<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Microbiol.</journal-id>
<journal-title>Frontiers in Microbiology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Microbiol.</abbrev-journal-title>
<issn pub-type="epub">1664-302X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fmicb.2025.1603255</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Microbiology</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>MGV-seq: a sensitive and culture-independent method for detecting microbial genetic variation</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Li</surname> <given-names>Lun</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x2020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Kong</surname> <given-names>Weiyu</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x2020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2993657/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Sun</surname> <given-names>Jing</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Jiang</surname> <given-names>Yongzhong</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Li</surname> <given-names>Tiantian</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Xia</surname> <given-names>Zhihui</given-names></name>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Zhou</surname> <given-names>Junfei</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Fang</surname> <given-names>Zhiwei</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Chen</surname> <given-names>Lihong</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/477066/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Feng</surname> <given-names>Shun</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Song</surname> <given-names>Huiyin</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2791803/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Xiao</surname> <given-names>Huafeng</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Zhang</surname> <given-names>Baolong</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2911674/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Fang</surname> <given-names>Bin</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/931050/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Peng</surname> <given-names>Hai</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1306712/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Gao</surname> <given-names>Lifen</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
<xref ref-type="corresp" rid="c002"><sup>&#x002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1743231/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Institute for Systems Biology, Jianghan University, Wuhan</institution>, <addr-line>Hubei</addr-line>, <country>China</country></aff>
<aff id="aff2"><sup>2</sup><institution>Hubei Provincial Center for Disease Control and Prevention, Wuhan</institution>, <addr-line>Hubei</addr-line>, <country>China</country></aff>
<aff id="aff3"><sup>3</sup><institution>Hainan Key Laboratory for Sustainable Utilization of Tropical Bioresources, College of Tropical Crops, Hainan University, Haikou</institution>, <addr-line>Hainan</addr-line>, <country>China</country></aff>
<aff id="aff4"><sup>4</sup><institution>Mingliao Biotechnology Co., Ltd., Wuhan</institution>, <addr-line>Hubei</addr-line>, <country>China</country></aff>
<aff id="aff5"><sup>5</sup><institution>Wuhan Zhongwei Gene Technology Co., Ltd., Wuhan</institution>, <addr-line>Hubei</addr-line>, <country>China</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: James T. Tambong, Agriculture and Agri-Food Canada (AAFC), Canada</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Bhabesh Dutta, University of Georgia, United States</p><p>Marcus Dillon, University of Toronto Mississauga, Canada</p></fn>
<corresp id="c001">&#x002A;Correspondence: Hai Peng, <email>penghai138@163.com</email></corresp>
<corresp id="c002">Lifen Gao, <email>lfgao@jhun.edu.cn</email></corresp>
<fn fn-type="equal" id="fn002"><p><sup>&#x2020;</sup>These authors have contributed equally to this work and share first authorship</p></fn>
</author-notes>
<pub-date pub-type="epub">
<day>25</day>
<month>06</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>16</volume>
<elocation-id>1603255</elocation-id>
<history>
<date date-type="received">
<day>31</day>
<month>03</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>06</day>
<month>06</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2025 Li, Kong, Sun, Jiang, Li, Xia, Zhou, Fang, Chen, Feng, Song, Xiao, Zhang, Fang, Peng and Gao.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Li, Kong, Sun, Jiang, Li, Xia, Zhou, Fang, Chen, Feng, Song, Xiao, Zhang, Fang, Peng and Gao</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<sec>
<title>Background</title>
<p>Precise detection of microbial genetic variation (MGV) at the strain level is essential for reliable disease diagnosis, pathogen surveillance, and reproducible research. Current methods, however, are constrained by limited sensitivity, specificity, and dependence on culturing. To address these challenges, we developed MGV-Seq, an innovative culture-independent approach that integrates multiplex PCR, high-throughput sequencing, and bioinformatics to analyze multiple dispersed nucleotide polymorphism (MNP) markers, enabling high-resolution strain differentiation.</p>
</sec>
<sec>
<title>Methods</title>
<p>Using <italic>Xanthomonas oryzae</italic> as a model organism, we designed 213 MNP markers derived from 458 genome assemblies. Method validation encompassed reproducibility, accuracy, sensitivity (detection limit), and specificity using laboratory-adapted strains, artificial DNA mixtures, and uncultured rice leaf samples. Performance was benchmarked against whole-genome sequencing (WGS) and LoFreq variant calling.</p>
</sec>
<sec>
<title>Results</title>
<p>MGV-Seq achieved 100% reproducibility and accuracy in major allele detection, with sensitivity down to 0.1% (<italic>n</italic> = 12 strains) for low-abundance variants and significantly higher specificity than LoFreq. Analysis to 40 <italic>X. oryzae</italic> strains revealed widespread heterogeneity (90% of strains) and misidentification (e.g., HN-P5 as <italic>Xoc</italic>). Homonymous strains exhibited significant genetic and phenotypic divergence, attributed to contamination rather than mutation. MGV-Seq successfully identified dominant strains and low-frequency variants in rice leaf samples and authenticated single-colony strains with 100% major allele similarity.</p>
</sec>
<sec>
<title>Conclusion</title>
<p>MGV-Seq establishes a robust, high-throughput solution for strain identification, microevolution monitoring, and authentication, overcoming limitations of culture-dependent and metagenomics-based methods. Its applicability extends to other microorganisms, offering potential for clinical, agricultural, and forensic diagnostics.</p>
</sec>
</abstract>
<kwd-group>
<kwd>microbial genetic variation</kwd>
<kwd>multiple dispersed nucleotide polymorphism markers</kwd>
<kwd>multiplex polymerase chain reaction</kwd>
<kwd>MGV-Seq</kwd>
<kwd>strain purification</kwd>
<kwd>identification authentication</kwd>
</kwd-group>
<counts>
<fig-count count="6"/>
<table-count count="1"/>
<equation-count count="0"/>
<ref-count count="43"/>
<page-count count="16"/>
<word-count count="10018"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Microbe and Virus Interactions with Plants</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="S1" sec-type="intro">
<title>1 Introduction</title>
<p>Plant pathogens severely damage crops, vegetables, and fruits. Based on the infected hosts and pathogenic traits, plant pathogenic species can be classified into subspecies, pathovar, subpathovar, and strains, representing considerable variability in their pathogenicity (<xref ref-type="bibr" rid="B34">Rademaker et al., 2005</xref>). Therefore, accurate and sensitive detection of taxonomic variations and accurate identification of phytopathogenic microorganisms at the strain level in unprocessed host samples are crucial for microbe-related disease diagnosis and prevention. However, the genetic and phenotypic similarity of taxa beyond species level challenges isolation and culturing, limiting the effectiveness of conventional biochemical, pathogenicity and serological analyses (<xref ref-type="bibr" rid="B4">Catara et al., 2021</xref>).</p>
<p>Identification of microbial strains is also essential for plant research involving phytopathogenic microorganisms. A well-curated strain is generally regarded as a genetically homogeneous population. However, laboratory-adapted microbial strains are usually used for laboratory-based research for a long time or distributed to laboratories worldwide. All the issues associated with cell lines&#x2014;spontaneous mutation, contamination, and mislabeling&#x2014;are likely to appear in phytopathogenic microorganisms during culture, storage, propagation, and passage (<xref ref-type="bibr" rid="B1">Andreu and Gibert, 2008</xref>, <xref ref-type="bibr" rid="B8">Feng et al., 2023</xref>, <xref ref-type="bibr" rid="B43">Zeng et al., 2023</xref>). The heterogeneity of cell lines has led to incomparable, inconsistent, or even contradictory results (<xref ref-type="bibr" rid="B15">Huang et al., 2017</xref>, <xref ref-type="bibr" rid="B21">Liu et al., 2019</xref>, <xref ref-type="bibr" rid="B22">Lorsch et al., 2014</xref>, <xref ref-type="bibr" rid="B26">Masters, 2012</xref>, <xref ref-type="bibr" rid="B43">Zeng et al., 2023</xref>). The importance of cell line authentication has attracted a great deal of attention lately. However, these critical issues have rarely been considered in studies and applications using phytopathogenic microorganisms. Technically, addressing this problem is also a challenge to existing approaches.</p>
<p>The existing approaches for detecting genetic variations within and across microorganisms can be divided into marker-dependent and marker-free categories. Marker-dependent approaches differentiate strains based on their genetic marker genotypes. Based on how markers are genotyped, they can be further categorized by length-based methods, such as multilocus variable-number tandem repeat analysis and pulsed-field gel electrophoresis analysis (<xref ref-type="bibr" rid="B12">Grissa et al., 2008</xref>, <xref ref-type="bibr" rid="B23">Machado et al., 2014</xref>, <xref ref-type="bibr" rid="B30">Poulin et al., 2015</xref>), and sequence-based methods, including multilocus sequence typing (MLST), DNA barcoding method such as 16S rRNA gene (16S) and internal transcribed spacer (ITS) (<xref ref-type="bibr" rid="B2">Antil et al., 2023</xref>), as well as whole-genome sequencing-single-nucleotide polymorphism (WGS-SNP) (<xref ref-type="bibr" rid="B23">Machado et al., 2014</xref>, <xref ref-type="bibr" rid="B24">Maiden, 2006</xref>, <xref ref-type="bibr" rid="B35">Roychowdhury et al., 2019</xref>, <xref ref-type="bibr" rid="B36">Saltykova et al., 2018</xref>). Length-based methods require minimal equipment and are convenient to use. However, they can only identify obvious differences in the length of the tested markers; they fail to detect nucleotide modifications and are dependent on the simultaneous experimentation of standard samples.</p>
<p>In contrast, the sequence-based MLST and DNA barcoding methods both utilize PCR amplification followed by targeted sequencing to examine nucleotide variations in a limited set of genetic markers. The MLST method focuses on a small number of species-specific markers (typically &#x003C;10 loci) in several evolutionarily conserved housekeeping genes (<xref ref-type="bibr" rid="B14">Hajri et al., 2012</xref>, <xref ref-type="bibr" rid="B24">Maiden, 2006</xref>), and thus is only suitable for long-term monitoring and cannot track variations that occur over a short period. DNA barcoding method targets one or two fast-evolving regions within conserved genes (e.g., 16S, ITS), and are widely employed for cross-species identification. Nevertheless, this method is limited in distinguishing closely related taxa or strains due to insufficient sequence divergence. Furthermore, its reliance on minimal genetic markers makes results susceptible to PCR amplification failures caused by primer-template mismatches.</p>
<p>On the otherhand, the WGS-SNP technology can detect thousands of SNPs across genomes of microbes (<xref ref-type="bibr" rid="B34">Rademaker et al., 2005</xref>). However, SNPs are typically biallelic markers; therefore, they are inadequate to represent the underlying allelic diversity of a microbial population. Furthermore, SNP genotyping error rates are estimated to be as high as 0.9&#x2013;7.9% (<xref ref-type="bibr" rid="B3">Bayer et al., 2017</xref>, <xref ref-type="bibr" rid="B6">Fang et al., 2017</xref>, <xref ref-type="bibr" rid="B19">Li, 2016</xref>, <xref ref-type="bibr" rid="B38">Sasaki et al., 2018</xref>), making it difficult to accurately detect variations in frequencies lower than the technical error rates. Consequently, WGS-SNP technology remains largely restricted to applications involving homogeneous, isolated strains. However, only 0.1&#x2013;1% of microorganisms have been reported to be cultivated (<xref ref-type="bibr" rid="B39">Solden et al., 2016</xref>). Coupled with the bottleneck of culture isolation, culture-based approaches often result in biased selection of the dominant pathogenic species or genotype. In addition, they are time-consuming, ineffective, and unreliable for identifying strains or detecting sample variations.</p>
<p>In contrast, metagenomic next-generation sequencing (mNGS) is not limited to molecular markers across genomes and is usually applied to characterize uncultured microbial communities. However, mNGS requires ultra-large data sequencing, anywhere from 7.98 to 18.00 Gbp (<xref ref-type="bibr" rid="B27">Newberry et al., 2020</xref>), and intensive data computing to accurately distinguish the true targets from the numerous common epiphytic and entophytic bacterial species associated with plants. Furthermore, the background bacteria are often found in pathogen testing laboratories (<xref ref-type="bibr" rid="B27">Newberry et al., 2020</xref>), resulting in limited resolution at the species level and a limited capability to detect variations in taxon units beyond species (<xref ref-type="bibr" rid="B16">Jongman et al., 2020</xref>).</p>
<p>For accurate, efficient, and sensitive detection of microbial genetic variation (MGV), we developed MGV-Seq, a culture-independent MGV detection method. MGV-Seq combines multiplex polymerase chain reaction (PCR) amplification profiling, high-throughput sequencing, and a customized computational pipeline to capture all variations in hundreds of multiple dispersed nucleotide polymorphism (MNP) markers. MNP markers are a novel type of DNA marker consisting of multiple SNPs dispersed within short genomic segments, i.e., 250 nucleotides (<xref ref-type="bibr" rid="B7">Fang et al., 2021</xref>, <xref ref-type="bibr" rid="B10">Gao et al., 2023</xref>). Crucially, since this length falls within the sequencing read length, the haplotypes (i.e., the specific combination of multiple SNPs) can be directly derived from individual sequencing reads. Theoretically, the diversity of haplotypes increases exponentially with the number of SNPs covered, providing greater higher discriminative power for capturing diverse alleles within complex microbial communities. Coupled with multiplex PCR amplification of hundreds of genome-wide MNP markers and high-throughput sequencing, MGV-Seq is capable of simultaneous acquisition of high-resolution genetic profiles for the target microbe in an uncultured sample, avoiding time-consuming culture isolation processes. Furthermore, our integration of a statistical model to accurately capture alleles across varying abundance levels from sequencing data offers an ideal approach for strain identification, strain authentication, and monitoring strain microevolution.</p>
<p>In this study, we used the top bacterial plant pathogen <italic>Xanthomonas oryzae</italic> (<xref ref-type="bibr" rid="B25">Mansfield et al., 2012</xref>), which contains two rice pathovars, <italic>oryzae</italic> (<italic>Xoo</italic>) and <italic>oryzicola</italic> (<italic>Xoc</italic>), as a model microorganism (<xref ref-type="bibr" rid="B28">Ni&#x00F1;o-Liu et al., 2006</xref>). We first proved the high reproducibility and accuracy of MGV-Seq for identifying pathogens at the strain level in well-curated <italic>X. oryzae</italic> strains, as well as its high sensitivity and specificity for detecting low-abundance variations in artificial mixtures of strain DNA. We then applied MGV-Seq to detect the widespread genetic variations within laboratory strains and the microevolution between homonymous strains from the same parent, after which we sensitively identified <italic>X. oryzae</italic> strains and low-abundance variations directly from uncultured rice leaf samples. Moreover, we demonstrated the feasibility of MGV-Seq-based strain authentication to ensure the genetic similarity between daughter lines and the standard strain. Although MGV-Seq was used to analyze one phytopathogenic microorganism, we propose its broad applicability to other microorganisms.</p>
</sec>
<sec id="S2" sec-type="materials|methods">
<title>2 Materials and methods</title>
<sec id="S2.SS1">
<title>2.1 DNA extraction of <italic>X. oryzae</italic> strains</title>
<p>The P1&#x2013;P10 strains from the International Rice Research Institute (IRRI) and the daughter lines from the Beijing, Wuhan, and Hainan laboratories used in this study are referred to as IRRI-P1&#x2013;IRRI-P10, BJ-P1&#x2013;BJ-P10, WH-P1&#x2013;WH-P10, and HN-P1&#x2013;HN-P10, respectively (<xref ref-type="supplementary-material" rid="FS1">Supplementary Figure S1</xref>). DNA of the IRRI-P1 to IRRI-P10 strains was provided by the IRRI laboratory. For DNA preparation of strains BJ-P1 to BJ-P10, WH-P1 to WH-P10, and HN-P1 to HN-P10, the cryopreserved bacterial cultures were reactivated via streaking on Potato-Sugar-Agar (PSA) medium without single-colony purification in their respective storage laboratories. After 72 h of incubation at 28&#x00B0;C, 100 mL of the activated inoculum was collected for DNA extraction using Rapid Bacterial Genomic DNA Isolation Kit (B518225, Sangon Biotech, China). For DNA preparation of rice leaf samples, rice leaves exhibiting bacterial blight symptoms were collected from the same paddy field area. Approximately 3 cm of tissue from the junction between diseased and healthy regions was excised. Ten leaves were pooled, and their surfaces were first cleaned with sterile water. Subsequently, DNA was extracted from the tissue powder using the aforementioned rapid bacterial genomic DNA extraction kit. DNA quantity was measured using Qubit Fluorometric Quantitation (Thermo Fisher Scientific, MA, United States).</p>
</sec>
<sec id="S2.SS2">
<title>2.2 Virulence evaluation of Xoo strains</title>
<p>The cryopreserved <italic>Xoo</italic> strains were reactivated by inoculation on Potato-Sucrose-Agar (PSA) medium and cultured at 28&#x00B0;C for 72 h. The activated bacterial cells were suspended in sterile water to prepare a bacterial suspension with a concentration of approximately 10<sup>9</sup> CFU/mL. At the rice tillering stage, inoculation was performed using the leaf-clipping method on 4&#x2013;5 fully expanded leaves per plant. Lesion length (LL) was measured on 10&#x2013;15 infected leaves from five biological replicates 12 days post-inoculation. LL &#x2265;10 cm, 10&#x003E;LL &#x2265;5 cm, and LL &#x003C;5 cm were classified as susceptible, moderately resistant, and resistant, respectively, based on the disease rating system for LL (<xref ref-type="bibr" rid="B32">Quibod et al., 2020</xref>). We performed two-tailed Student&#x2019;s <italic>t</italic>-tests to compare lesion lengths between experimental groups, applying a stringent significance threshold (<italic>p</italic> &#x003C; 0.005) with Bonferroni correction for multiple comparisons when applicable.</p>
</sec>
<sec id="S2.SS3">
<title>2.3 Design of MNP markers and multiplex PCR primers</title>
<p>The design of the MNP markers and multiplex PCR primers for <italic>Xanthomonas oryzae</italic> followed our previous report with minor adjustments (<xref ref-type="bibr" rid="B7">Fang et al., 2021</xref>, <xref ref-type="bibr" rid="B10">Gao et al., 2023</xref>). The P6 (PXO99A) genome (<xref ref-type="bibr" rid="B37">Salzberg et al., 2008</xref>) was corrected using the genome data of WH-P6 and used as the synthetic reference genome to call SNPs in 458 <italic>X. oryzae</italic> strains using Bowtie2 (version 2.1.0) (<xref ref-type="bibr" rid="B17">Langmead and Salzberg, 2012</xref>). A sliding window of 275 base pairs (bp) was used for genome scanning with a 5 bp increment for the MNP loci design. The windows with discriminative power, which were defined as <italic>t/c(N,2)</italic>, where <italic>c(N,2)</italic> is the number of pairs among <italic>N</italic> strains used, and <italic>t</italic> is the number of pairs, each of which had at least two dispersed SNPs within the window, &#x2265;0.1, were chosen for the multiplex PCR primer design. Primers were designed using Ion AmpliSeq Designer<sup><xref ref-type="fn" rid="footnote1">1</xref></sup> and synthesized by Thermo Fisher Scientific (Waltham, MA, United States).</p>
</sec>
<sec id="S2.SS4">
<title>2.4 Library construction and high-throughput sequencing</title>
<p>For all the samples in this study (including <italic>Xoo</italic> strains and rice leaf samples), sequencing libraries were constructed from 10 ng genomic DNA per sample using the Ion AmpliSeq Library Kit 2 (Cat# 4475345, Thermo Fisher Scientific, United States). The constructed amplicon library was quantified using the TaqMan probe method and mixed in equimolar amounts for sequencing on an Ion S5 sequencer (A27212, Thermo Fisher Scientific, Waltham, MA, United States) by single-end sequencing with a 300-bp read length.</p>
</sec>
<sec id="S2.SS5">
<title>2.5 Estimation of the reproducibility and accuracy of MGV-seq</title>
<p>Three technical repeats and two sequencing repeats of 12 strains (WH-P1, BJ-P1, HN-P1, WH-P2, BJ-P2, HN-P2, WH-P9, BJ-P9, HN-P9, WH-P10, BJ-P10, and HN-P10) were performed to estimate the accuracy and reproducibility of MGV-Seq. At the first sequencing batch, each of the 12 strains was prepared in three biological replicates (designated as rep-1, rep-2, and rep-3), thus a total of 12&#x00D7;3 = 36 technical repeats of 12 strains were generated. At the second sequencing batch, an additional replicate for the same 12 strains was performed (designated as rep-4). Consequently, this design yielded 3 batch replicates per strain (rep-1 vs. rep-4, rep-2 vs. rep-4, rep-3 vs. rep-4), resulting in a total of 36 batch replicates across all 12 strains (12 strains &#x00D7; 3 comparisons each). If the major allele of one MNP locus between two repeats is identical, the MNP locus is reproducible. The reproducibility of MGV-Seq was calculated as the ratio of the number of reproducible pairs to all genotype pairs. The reproducible major alleles were considered correct because the chance that the major alleles of one MNP locus detected by two repeats were consistent but incorrect should be relatively small. Therefore, the accuracy of MGV-Seq for detecting major alleles was 0.5+0.5&#x002A; reproducibility.</p>
</sec>
<sec id="S2.SS6">
<title>2.6 A containerized computational tool for MGV-Seq</title>
<p>The tool containerized the computation via the following steps:</p>
<sec id="S2.SS6.SSS1">
<title>2.6.1 MNP allele typing</title>
<p>FASTX-Toolkit (version 0.0.14)<sup><xref ref-type="fn" rid="footnote2">2</xref></sup> was used to pre-process the raw reads. First, base pairs with low-quality scores (&#x003C;20) were trimmed, and the resultant reads &#x003C;50 bp were removed. Then, the processed reads were mapped to a synthetic reference (see Design of MNP marker and multiplex PCR primer subsection) using Bowtie2 (version 2.1.0).</p>
<p>For each read that covered the entire genomic region of an MNP locus, the sequence located within the locus was used as its allele. To reduce errors caused by incorrect read alignment, any variation within consecutive mismatches with reference (mismatch interval &#x2264;2 bp), repeat sequences (SSR), indels, and their flanking 2 bp of an allele were ignored. In addition, loci with at least one typed read were considered detected loci, and those with &#x003E;100 typed reads were deemed valid MNP loci.</p>
</sec>
<sec id="S2.SS6.SSS2">
<title>2.6.2 MNP genotyping</title>
<p>Valid MNP loci were genotyped (<xref ref-type="supplementary-material" rid="FS2">Supplementary Figure S2</xref>). Briefly, all alleles identified in a valid MNP locus were referred to as candidate alleles of the locus in a microorganism strain. The candidate allele with the most supported reads was the major allele of the MNP locus in the strain if it was supported by at least 50% of the mapped reads on the locus; otherwise, this locus was excluded from the following analysis.</p>
<p>The remaining candidate alleles were a mixture of genuine minor alleles at lower frequencies and artifacts caused by PCR or sequencing errors. To distinguish genuine alleles from false ones, candidate alleles were first subjected to strand bias filters: (1) if the strand bias was greater than 10-fold or the differences between the strand bias of the candidate allele or that of the major allele were more than five-fold, the candidate allele was removed; (2) Fisher&#x2019;s exact tests were applied to determine the strand bias, and candidate alleles with multiple tests corrected <italic>p</italic>-values &#x003C;0.05 were abandoned.</p>
<p>Subsequently, a statistical model was applied. A candidate allele with <italic>c</italic> supporting reads was primarily assumed to be the product of the major allele solely due to amplicon sequencing errors. Under this hypothesis, the number of supporting reads of a candidate allele follows a binomial distribution with parameters <italic>d</italic> (marker sequencing depth) and <italic>e</italic> (amplicon sequencing error rate). If the candidate allele had an excessively larger number of supporting reads than expected, namely, if P(<italic>X</italic>&#x2265;<italic>c</italic>) was significantly low, the null hypothesis was overridden. When multiple candidate alleles were tested, the <italic>p</italic>-values were corrected for multiple testing, and candidate alleles with an FDR &#x003C;0.5% were deemed true minor alleles.</p>
</sec>
<sec id="S2.SS6.SSS3">
<title>2.6.3 Estimation of genotyping parameters of MGV-Seq</title>
<p>Amplicon sequencing errors supposedly occur independently across reads; thus, the probability of multiple errors occurring is lower within a read sequence than that of a single error. Therefore, the parameter <italic>e</italic> used in the statistical model depends on the number of SNPs (<italic>n</italic>) between the candidate allele and major allele. The minor alleles detected at the homozygous MNP loci can be regarded as products of technical error. In this study, the parameter <italic>e</italic> for each <italic>n</italic>, denoted as <italic>e</italic> (<italic>n</italic>), had the highest ratio of reads that an erroneous allele bearing <italic>n</italic> SNPs could take up at a homogeneous MNP locus. Microorganisms typically live in populations, and monoclonal organisms are often composed of individuals of multiple generations. Therefore, the homozygous MNP loci in microorganisms could not be validated. Note that rice is a diploid plant, and F<sub>1</sub> hybrids are produced by hybridization. The MNP loci identified as homozygous in both parents of hybrids are expected to be homozygous in the hybrid. Any minor alleles identified at these MNP loci can be attributed to amplicon sequencing errors. Therefore, we utilized the minor alleles detected at the expected homozygous MNP loci in the individual leaves of rice plants by amplicon sequencing to estimate the parameters [<italic>e</italic> (<italic>n</italic>1)] and MNPs [<italic>e</italic> (<italic>n</italic>1)]. To minimize sampling errors, only loci with depth &#x003E;1,000 were considered. Based on the 930 MNP markers of the 16 individual rice plants, <italic>e</italic> (<italic>n</italic> = 1) and <italic>e</italic> (<italic>n</italic> &#x003E; 1) were estimated to be 1.03 and 0.0994%, respectively.</p>
</sec>
</sec>
<sec id="S2.SS7">
<title>2.7 Minor alleles detected by LoFreq</title>
<p>As reported previously (<xref ref-type="bibr" rid="B42">Wilm et al., 2012</xref>), a two-step strategy was used to call minor alleles using LoFreq. We first obtained consensus sequences from the original mapping results and used them as a reference in the second mapping round. Then, the final mapping results were analyzed using LoFreq (version 2.1.4).</p>
</sec>
<sec id="S2.SS8">
<title>2.8 Strain identification and authentication</title>
<p>The MNP genotypes of a tested strain or sample were compared with those of the strain to be compared using in-house scripts to obtain the major allele similarity (MAS) between the two, i.e., the ratio of the number of MNP loci with the same major alleles to the number of MNP loci detected commonly. For strain identification, the MNP genotypes of the tested samples were compared with those of the conspecific strains. The conspecific strain in the comparison group with the highest MAS was identified as the dominant strain in the tested sample. For strain authentication, the MNP genotypes of the tested strains were compared with those of the standard strains. The MAS threshold for strain authentication was established based on experimental requirements. In this study, we used a very stringent criterion, i.e., we considered that the strains were adequately authenticated when they had 100% MAS.</p>
</sec>
<sec id="S2.SS9">
<title>2.9 WGS</title>
<p>Whole genomes of IRRI-P1&#x2013;IRRI-P10 were sequenced at Molbreeding Biotechnology Co., Ltd (Shijiazhuang, Hebei, China) on the Illumina HiSeq X Ten platform (Illumina, Inc., San Diego, CA, United States) with 150 bp read length. In addition, DNA libraries were constructed using the DNA library kit V3.1 (Molbreeding Biotechnology Co., Ltd.).</p>
</sec>
<sec id="S2.SS10">
<title>2.10 Sanger sequencing validation of MNP genotypes</title>
<p>Genotypes of the MNP loci among samples were validated using the same primers used in MGV-Seq to amplify the MNP loci, and Sanger sequencing was used for the amplification products (TsingKe Company, Wuhan, Hubei, China).</p>
</sec>
<sec id="S2.SS11">
<title>2.11 WGS alleles and published datasets</title>
<p>The WGS reads were first mapped to the reference genome using Bowtie2 (version 2.1.0) (<xref ref-type="bibr" rid="B17">Langmead and Salzberg, 2012</xref>) and consensus sequences within the MNP loci were considered major alleles. Nucleotides with &#x003C;20 reads were masked to mitigate errors caused by low sequencing depths. Furthermore, the reference sequence of the MNP loci was compared with published assemblies using BLAST (version 2.2.26), and the most similar sequences were taken as the major alleles. Finally, the minor alleles detected using WGS were called by LoFreq using the downloaded complete genomes as respective references (see &#x201C;Minor alleles detected by LoFreq&#x201D; subsection).</p>
</sec>
<sec id="S2.SS12">
<title>2.12 Statistical analysis</title>
<p>All statistical analyses were conducted in R (version 3.5.1). Student&#x2019;s <italic>t</italic>-tests and Fisher&#x2019;s exact tests were performed by &#x201C;t.test&#x201D; function and &#x201C;fisher.test&#x201D; functions, respectively. Cluster analysis based on MNP genotypes was carried out as follows. First, pairwise genetic distances between strains were calculated as the complement of MAS (defined in section 2.8). On the basis of the calculated distance matrix, hierarchical clustering was then carried out using R&#x2019;s &#x201C;hclust&#x201D; function. Finally, the resulting cluster dendrogram was visualized using the &#x201C;ggtree&#x201D; package (version 1.12.7).</p>
</sec>
</sec>
<sec id="S3" sec-type="results">
<title>3 Results</title>
<sec id="S3.SS1">
<title>3.1 Design and characterization of MGV-Seq</title>
<sec id="S3.SS1.SSS1">
<title>3.1.1 MGV-Seq&#x2014;the method</title>
<p>We designed the MGV-Seq method to profile MGVs. The MGV-Seq procedure began with the selection of MNP markers (<xref ref-type="bibr" rid="B7">Fang et al., 2021</xref>) using the (re)sequencing data of representative strains from a species (see Materials and methods section). Briefly, the genome database of target species was constructed by collecting genome assemblies of conspecific strains from a public database, such as National Center for Biotechnology Information. Then, the genome assembly of a representative strain was used as the reference genome to call SNPs in the other genome assemblies, and a sliding window was used to scan those genomes assemblies to screen windows with at least two SNPs. The length of the sliding window and step between windows is determined based on the length of the amplicon that can be processed by the sequencer used. The windows with high discriminative power within species and uniform distribution in the reference genome were selected as candidates of MNP markers (see Materials and methods section). Primers for the multiplex amplification of the candidate MNP markers were designed using the available tools and synthesized after adding adaptors adapted to the sequencer used. Using the synthesized primers and a suitable kit for amplicon sequencing, a small number of representative strains was tested to screen MNP loci that could be detected in all strains to form the final MNP marker set and primer pools of MGV-Seq.</p>
<p>Next, MGV-Seq can be used to detect MNP markers in samples from pure culture to uncultured diseased host samples (<xref ref-type="fig" rid="F1">Figures 1A,B</xref>). The PCR amplification products of each sample were ligated with a unique DNA barcode. Finally, the barcoded PCR products were mixed to form a library for high-throughput multiplex sequencing (<xref ref-type="fig" rid="F1">Figure 1C</xref>). A customized computational tool was developed for MNP genotyping (<xref ref-type="fig" rid="F1">Figure 1C</xref>). Briefly, we first assigned sequencing reads to samples based on DNA barcodes and mapped the reads to the reference genome. Next, all reads of an MNP locus were tallied to call for candidate alleles. The candidate allele supported by &#x003E;50% of the total reads on a locus was defined as the major allele of the locus. All candidate minor alleles for an MNP locus were assumed to originate from the major allele due to amplicon sequencing errors. The true minor alleles of the locus were determined by strand bias filtering and statistical modeling (<xref ref-type="fig" rid="F1">Figure 1C</xref>, see Materials and methods section). Finally, the MNP genotypes of each sample were constructed by integrating the alleles of all detected MNP loci.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption><p>Workflow of MGV-Seq. <bold>(A)</bold> The design of MNP markers includes screening polymorphic candidate MNP loci from multiple genomes of the target microorganism, designing multiplex PCR primers for these loci, and validating and determining the valid MNP loci. &#x002A; with different colors represents A, T, C, and G bases. Arrows of the same color in different directions represent a pair of primers. <bold>(B)</bold> Samples that MGV-Seq can detect. <bold>(C)</bold> The main steps of MGV-Seq include library preparation, high-throughput sequencing of the libraries, and bioinformatics analysis of the sequencing data. The red &#x002A; represent the SNP variations among reads. <bold>(D)</bold> MGV-Seq-based procedure for heterogeneous microbial strain determination, strain authentication, purification monitoring, and identification. The arrows in the figure of two different colors represent two application procedures of MGV-Seq. The orange arrows correspond to homogeneous sample processing, while the blue arrows represent the experimental processes for heterogeneous samples. MNP, multiple dispersed nucleotide polymorphism; PCR, polymerase chain reaction; MGV, microbial genetic variation.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmicb-16-1603255-g001.tif"/>
</fig>
<p>A procedure for applying MGV-Seq to homogeneous and heterogeneous strain determination, strain identification, strain authentication, and monitoring strain microevolution was proposed (<xref ref-type="fig" rid="F1">Figure 1D</xref>). The heterogeneity of the tested samples was determined based on the number of minor alleles detected in the MNP genotypes. By comparing the MNP genotypes of the tested sample with those of the conspecific, parental, or standard strains, MGV-Seq can be further used for identifying the dominant strain in the tested sample, monitoring the microevolution of daughter strain lines, and authenticating the strain to be used (see Materials and methods section).</p>
</sec>
<sec id="S3.SS1.SSS2">
<title>3.1.2 Reproducibility, accuracy, sensitivity, and specificity of MGV-Seq</title>
<p>This study used the phytopathogen <italic>Xanthomonas oryzae</italic> (<italic>X. oryzae</italic>) as a model microorganism. A set of 213 MNP markers was developed (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S1</xref>; see Materials and methods section) based on 458 genome assemblies of <italic>X. oryzae</italic> strains downloaded from the National Center for Biotechnology Information (NCBI), including 427 <italic>Xoo</italic> and 19 <italic>Xoc</italic> genomes (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S2</xref>). Widely used P1&#x2013;P10 <italic>Xoo</italic> strains from IRRI (<xref ref-type="bibr" rid="B33">Quibod et al., 2016</xref>) and daughter lines from three different laboratories in Beijing, Wuhan, and Hainan, China, were collected (<xref ref-type="supplementary-material" rid="FS1">Supplementary Figure S1</xref>). The P1&#x2013;P10 strains were short-handed for PXO61, PXO86, PXO79, PXO71, PXO112, PXO99, PXO145, PXO280, PXO339, and PXO341 in this study. Biological samples with the same name but from different laboratories are referred to as homonyms in this study, while DNA samples sequenced at the same run are referred to as a sequencing batch. Repeats of the same DNA tested in the same and different sequencing batches are called technical and batch repeats, respectively.</p>
<p>The dominant strain within each sample was identified by comparing the major alleles on the detected MNP loci with the reference genotypes of the conspecific strains. Therefore, the reproducibility of MGV-Seq for strain identification can be estimated using the reproducibility of MGV-Seq for detecting major alleles. A total of 4,971 MNP pairs among the 36 technical repeats and 4,997 MNP pairs between the 36 batch repeats were generated. The major alleles in all MNP pairs were 100% reproducible (<xref ref-type="fig" rid="F2">Figure 2A</xref> and <xref ref-type="supplementary-material" rid="TS1">Supplementary Tables S3</xref>, <xref ref-type="supplementary-material" rid="TS1">S4</xref>). Reproducible major alleles were regarded as correct because there is only a tiny chance that major alleles of one MNP locus, detected by two repeats, were consistent but incorrect. Therefore, MGV-Seq demonstrated 100% reproducibility and accuracy for the technical and batch repeats for major allele detection. These 12 strains included three homonyms of the four strains. Pairwise comparison of homonymous strains was performed to examine the reproducibility of MGV-Seq for detecting differences across strain homonyms. A total <inline-formula><mml:math id="IEQ1"><mml:mrow><mml:mn>4</mml:mn><mml:mtext>&#x00A0;</mml:mtext><mml:mo>&#x00D7;</mml:mo><mml:mtext>&#x00A0;</mml:mtext><mml:mrow><mml:mo>(</mml:mo><mml:mrow><mml:mtable><mml:mtr><mml:mtd><mml:mn>3</mml:mn></mml:mtd></mml:mtr><mml:mtr><mml:mtd><mml:mn>2</mml:mn></mml:mtd></mml:mtr></mml:mtable></mml:mrow><mml:mo>)</mml:mo></mml:mrow><mml:mtext>&#x00A0;</mml:mtext><mml:mo>=</mml:mo><mml:mtext>&#x00A0;</mml:mtext><mml:mn>12</mml:mn></mml:mrow></mml:math></inline-formula> pairs of homonymous strains were formed and each pair performed 4&#x00D7;4 = 16 repeats from the four replicates of each strain across the two batches. The reproducibility of MGV-Seq for detecting differences across strain homonyms was 100% (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S5</xref>).</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption><p>Evaluation of reproducibility, accuracy, sensitivity, and specificity of MGV-Seq. <bold>(A)</bold> The reproducibility and accuracy of MGV-Seq for detecting major alleles. <bold>(B)</bold> Sensitivity and <bold>(C)</bold> specificity of MGV-Seq for detecting low-abundance variant from artificial DNA mixtures of WH-P8 and WH-P6. BJ, Beijing; WH, Wuhan; HN, Hainan; Rep, replicate; MNP, multiple dispersed nucleotide polymorphism.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmicb-16-1603255-g002.tif"/>
</fig>
<p>The sensitivity and specificity of MGV-Seq were estimated based on its performance in detecting mixed samples. A mixed sample was characterized by the presence of minor alleles belonging to MNP genotypes of the different strains. MGV-Seq was tested on 14 single clones of WH-P8 designated as CP8-1&#x2013;CP8-14, eight single clones of WH-P6, designated as CP6-1&#x2013;CP6-8, and a series of mixtures by spiking the DNA of one WH-P8 clone into one WH-P6 clone at increasing reference ratios of 0.1, 0.3, 0.5, 0.7, 1, 3, 5, and 7% w/w, at the first sequencing batch (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S3</xref>). The genotyping results of the single clones and mixtures were analyzed and compared using the MGV-Seq-adapted computational pipeline developed by us and LoFreq, a previously reported ultra-sensitive SNP-based variation detecting tool (<xref ref-type="bibr" rid="B42">Wilm et al., 2012</xref>). Minor alleles detected in the mixture belonging to either WH-P6 minor alleles or WH-P8 major alleles were called true positive minor alleles (TP), while others were called false positive minor alleles (FP). The sensitivity and specificity of MGV-Seq were referred to as the lowest ratio of spiking DNA in the detected mixture and the percentage of TP to the sum of TP and FP, respectively. MGV-Seq reported 20&#x2013;44.4% TP for 0.1&#x2013;7% reference ratios of mixtures and achieved a perfect (100%) specificity (TP/(TP+FP)) at all mixture ratios (<xref ref-type="fig" rid="F2">Figures 2B,C</xref>). In comparison, the percentages of TP detected by LoFreq for the 0.1&#x2013;7% reference ratios ranged from 12&#x2013;51.7% (<xref ref-type="fig" rid="F2">Figure 2B</xref>). The results showed that the sensitivity of MGV-Seq and LoFreq in the detection of mixed samples was as low as 0.1%; however, LoFreq had lower specificity for six of the eight reference ratios tested, for example, 92.3% specificity at 0.5% mixture ratio (<xref ref-type="fig" rid="F2">Figure 2C</xref>).</p>
</sec>
<sec id="S3.SS1.SSS3">
<title>3.1.3 WGS versus MGV-Seq</title>
<p>For technical comparison, whole genomes of the 10 <italic>Xoo</italic> strains from IRRI were sequenced (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S6</xref>). A total of 1,400 MNP loci detected by MGV-Seq in the 10 <italic>Xoo</italic> strains (IRRI-P1-IRRI-P10 at the second sequencing batch) were covered by the WGS data, and the major alleles were 100% verified by the WGS data (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S4</xref>), suggesting the genotyping reliability of MGV-Seq.</p>
</sec>
</sec>
<sec id="S3.SS2">
<title>3.2 Applications of MGV-Seq</title>
<sec id="S3.SS2.SSS1">
<title>3.2.1 Homogeneity and heterogeneity of the laboratory-adapted strains</title>
<p>Genetic homogeneity or heterogeneity of the strain was determined based on the minor alleles detected at the MNP loci of each strain. We applied MGV-Seq to detect 40 laboratory-adapted <italic>Xoo</italic> strains corresponding to 10 distinct parental strains each with four homonyms at the second sequencing batch. The four homonymous strains were P1&#x2013;P10 from the IRRI and daughter lines preserved in the Beijing, Wuhan, and Hainan laboratories, designated as IRRI-P1&#x2013;IRRI-P10, BJ-P1&#x2013;BJ-P10, WH-P1&#x2013;WH-P10, and HN-P1&#x2013;HN-P10, respectively (<xref ref-type="supplementary-material" rid="FS1">Supplementary Figure S1</xref>). On average, 1,676,528 reads mapped to 203 MNP loci per strain, with an average coverage of 8,245 folds per marker (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S3</xref>). No minor alleles were detected at the MNP loci in four of the 40 <italic>Xoo</italic> strains, which were considered homogeneous strains. MGV-Seq identified 254 minor alleles in the remaining 36 <italic>Xoo</italic> strains (<xref ref-type="fig" rid="F3">Figure 3A</xref>), and the number of minor alleles in each strain ranged from 1&#x2013;37 (<xref ref-type="fig" rid="F3">Figure 3B</xref>), indicating that 90% of the tested laboratory-adapted strains were heterogeneous. Moreover, 31.9% (81 of 254) of minor alleles had a frequency &#x003C;0.9% (<xref ref-type="fig" rid="F3">Figure 3D</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S7</xref>), which was considered to be the lower limit of the SNP genotyping error rates (<xref ref-type="bibr" rid="B3">Bayer et al., 2017</xref>, <xref ref-type="bibr" rid="B6">Fang et al., 2017</xref>, <xref ref-type="bibr" rid="B19">Li, 2016</xref>, <xref ref-type="bibr" rid="B38">Sasaki et al., 2018</xref>). Therefore, accurately capturing low-frequency minor alleles using SNP genotyping-based methods is difficult.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption><p>Determination of genetic homogeneity and heterogeneity of the laboratory-adapted strains. <bold>(A,B)</bold> Number of minor alleles of the <italic>Xoo</italic> strains from different laboratories. <bold>(C)</bold> Frequency of minor alleles of the <italic>Xoo</italic> strains from different laboratories. <bold>(D)</bold> Number distribution of strains sharing the same minor alleles. BJ, Beijing; WH, Wuhan; HN, Hainan; IRRI, International Rice Research Institute; <italic>Xoo</italic>, <italic>Xanthomonas oryzae</italic> pv. <italic>oryzae</italic>.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmicb-16-1603255-g003.tif"/>
</fig>
</sec>
<sec id="S3.SS2.SSS2">
<title>3.2.2 Identification of laboratory-cultured strains</title>
<p>We then identified these 40 strains by conducting a genetic cluster analysis based on the MNP genotypes of the 40 strains, the previously published genomes of P1&#x2013;P10, named Pub-P1&#x2013;Pub-P10 (<xref ref-type="bibr" rid="B33">Quibod et al., 2016</xref>, <xref ref-type="bibr" rid="B37">Salzberg et al., 2008</xref>, <xref ref-type="bibr" rid="B40">Song et al., 2023</xref>). and the genome of BLS254, representative strain of <italic>Xoc</italic> (<xref ref-type="bibr" rid="B9">Feng et al., 2015</xref>, <xref ref-type="bibr" rid="B14">Hajri et al., 2012</xref>, <xref ref-type="bibr" rid="B18">Lee et al., 2005</xref>, <xref ref-type="bibr" rid="B33">Quibod et al., 2016</xref>, <xref ref-type="bibr" rid="B37">Salzberg et al., 2008</xref>). As illustrated by the cluster dendrogram, 39 of the 40 strains were <italic>Xoo</italic> pathovars, but one (HN-P5) was distant from all analyzed <italic>Xoo</italic> strains but closer with <italic>Xoc</italic> strains, which was identified as <italic>Xoc</italic> (<xref ref-type="fig" rid="F4">Figure 4A</xref>). Alarmingly, BJ-P2 and HN-P2 were not clustered with their parental strains but were closest to the P6 strains, WH-P3 and HN-P3 were closest to the P9 strains, WH-P4 and WH-P10 were closest to some P2 strains, and HN-P5 were not clustered with any <italic>Xoo</italic> strains, respectively, indicating that there was a deviation between the identified and claimed identities of these strains. Pairwise comparison of P1&#x2013;P10 homonymous strains from the four laboratories was performed to examine MGV across strain homonyms. Of the total <inline-formula><mml:math id="IEQ2"><mml:mrow><mml:mn>10</mml:mn><mml:mtext>&#x00A0;</mml:mtext><mml:mo>&#x00D7;</mml:mo><mml:mtext>&#x00A0;</mml:mtext><mml:mrow><mml:mo>(</mml:mo><mml:mrow><mml:mtable><mml:mtr><mml:mtd><mml:mn>4</mml:mn></mml:mtd></mml:mtr><mml:mtr><mml:mtd><mml:mn>2</mml:mn></mml:mtd></mml:mtr></mml:mtable></mml:mrow><mml:mo>)</mml:mo></mml:mrow><mml:mtext>&#x00A0;</mml:mtext><mml:mo>=</mml:mo><mml:mtext>&#x00A0;</mml:mtext><mml:mn>60</mml:mn></mml:mrow></mml:math></inline-formula> pairs, 17 (28.3%) had distinct major alleles with ratios of 14.0&#x2013;76.2% (<xref ref-type="fig" rid="F4">Figure 4B</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S8</xref>), suggesting that MGV was prevalent among these strain homonyms. An ultra-high ratio of major allele differences existed between HN-P5 and other P5 strains, confirming the MGV-Seq genotyping result that HN-P5 belonged to <italic>Xoc</italic> rather than <italic>Xoo</italic>.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption><p>Identification and detection of genetic difference of homonymous strains from different laboratories. <bold>(A)</bold> MNP genotypes-based cluster analysis of the <italic>Xoo</italic> strains identified in this study and strains with published genomes (Pub-P1&#x2013;Pub-P10). The published genomes, i.e., the four genomes of the P1 strains, were named Pub-P1-1, Pub-P1-2, Pub-P1-3, and Pub-P1-4. &#x002A;represents an unexpected clustering of the genome, and the genome of BLS256 as the representative strain of <italic>Xoc</italic>. <bold>(B)</bold> The ratio of distinct major alleles between the <italic>Xoo</italic> homonyms derived from different laboratories. <bold>(C)</bold> Distribution of the ratios of distinct major alleles of strains homonyms and the IRRI strains. <bold>(D)</bold> Validation of genotypes of MNP loci among homonymous strains. MNP, multiple dispersed nucleotide polymorphism; <italic>Xoo</italic>, <italic>Xanthomonas oryzae</italic> pv. <italic>oryzae</italic>. BJ, Beijing; WH, Wuhan; HN, Hainan; IRRI, International Rice Research Institute.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmicb-16-1603255-g004.tif"/>
</fig>
<p>The significant degree of genetic discrepancy in the reference strain across daughter lines raised concerns about the origin of the genetic differences. Thus, we analyzed the distribution of distinct major alleles among the daughter lines and different strains. Different <italic>Xoo</italic> strains diverged for much longer than daughter lines and thus should have accumulated more mutations. Distinct major alleles among the daughter lines (14.0&#x2013;76.2%, <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S8</xref>) were greater than those among different strains (0&#x2013;23.5%, <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S9</xref>), and the distribution of distinct major alleles in the daughter lines was discontinuous (<xref ref-type="fig" rid="F4">Figure 4C</xref>). We further analyzed the distribution of minor alleles among the 40 strains. Eleven (4.3%) of the 254 minor alleles appeared only once and differed from the major alleles at any locus, suggesting that spontaneous mutations may have caused them. The remaining 243 (95.7%) minor alleles were the same as the major alleles of the other 1&#x2013;39 <italic>Xoo</italic> strains tested together (<xref ref-type="fig" rid="F3">Figure 3C</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S7</xref>). These results strongly suggested that human error, such as sample mislabeling or contamination, rather than spontaneous mutations, might be the primary factor underlying the genetic discrepancy in the examined <italic>Xoo</italic> daughter lines captured by MGV-Seq. Genotypes of a randomly selected MNP locus in the strain homonyms were successfully validated using Sanger sequencing (<xref ref-type="fig" rid="F4">Figure 4D</xref>).</p>
<p>A large number of MGV of <italic>Xoo</italic> homonyms suggested that these strains might have distinct pathogenicity, one of the most important phenotypic features of <italic>Xoo</italic>. We investigated the pathogenicity of P2, P3, P4, P5, and P10 homonyms from Wuhan, Beijing, and Hainan on the susceptible rice variety IR24 and a near-isogenic line IRBB2 carrying the <italic>Xoo</italic> resistance gene <italic>Xa2</italic> (<xref ref-type="fig" rid="F5">Figure 5</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S10</xref>). Most of the 10 pairs of strain homonyms with distinct major alleles induced significantly different disease reactions in IR24 and IRBB2 rice varieties, showing the effect of MGV on pathogenicity reproducibility. One (WH-P3 vs. HN-P3) out of the five pairs of <italic>Xoo</italic> homonyms without distinct major alleles had a difference in lesion lengths (LL) of leaves, which may be due to the effect of minor alleles on pathogenicity reproducibility.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption><p>Pathogenicity differences among the <italic>Xanthomonas oryzae</italic> pv. <italic>oryzae</italic> strain homonyms. Names marked red indicates the strain with major alleles distinct from its homonyms. Green, yellow, and red circles represent disease reaction levels of susceptible, moderately resistant, and resistant, respectively. CP6-1, CP6-2, and CP6-6 are single colony-derived strains from WH-P6. BJ, Beijing; WH, Wuhan; HN, Hainan. Two-tailed Student&#x2019;s <italic>t</italic>-tests were performed to compare lesion lengths between homonymous strains. Homonymous strains with a <italic>p</italic>&#x003C; 0.005 were connected with black frame.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmicb-16-1603255-g005.tif"/>
</fig>
</sec>
<sec id="S3.SS2.SSS3">
<title>3.2.3 Identification of Xanthomas oryzae strains detected in rice leaf samples</title>
<p>We used MGV-Seq to detect <italic>X. oryzae</italic> strains in rice leaves&#x2014;with similar leaf blight symptoms&#x2014;collected from the same paddy field in 2019 and 2020. MGV-Seq reported 185 and 144 valid MNP loci in the two rice leaf samples, respectively. And 7 of the 185 valid MNP loci and 47 of the 144 valid MNP loci had minor alleles (<xref ref-type="fig" rid="F6">Figure 6A</xref>). The number of minor alleles of each MNP loci was one, with frequencies ranging from 0.5 to 19.8% (<xref ref-type="fig" rid="F6">Figures 6B,C</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S11</xref>), suggesting the presence of low-frequency genetic variation in these samples. We identified the strains in the two rice leaf samples by comparing the major alleles of MNP loci detected with the 458 genomes in the reference database of <italic>X. oryzae</italic>. To illustrate the genetic distance between the rice sample genotypes and the reference strains, we plotted a clustering tree (<xref ref-type="fig" rid="F6">Figure 6B</xref>). For the clarity and readability of the clustering tree, the 10 strains, with genomes closest to that of P1-P10, in the <italic>X. oryzae</italic> reference database based on the calculated average nucleotide identity, and the X-representative strains were selected as the reference strains for plotting the clustering tree. All the results presented in the heatmap (<xref ref-type="supplementary-material" rid="FS3">Supplementary Figure S3</xref>) and the clustering tree showed that the dominant strain in the 2019 sample was closest to the PXO211 strain of <italic>Xoo</italic>, whereas that in the 2020 sample was closest to the L8 strain of <italic>Xoc</italic> (<xref ref-type="fig" rid="F6">Figure 6B</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S12</xref>).</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption><p>Strain identification and variation detection in rice leaf samples. <bold>(A)</bold> Number of valid MNP loci, valid MNP loci with minor alleles and number of minor alleles detected in the DNA of rice leaf samples collected in 2019 and 2020. <bold>(B)</bold> Cluster analysis of the strains identified in rice leaf samples with the <italic>Xanthomonas oryzae</italic> pv. <italic>oryzae</italic> strains used in this study. <bold>(C)</bold> Positions and frequencies of the minor alleles detected in the DNA of the rice leaf samples collected in 2019 and 2020. The numbers in parentheses represent the base positions of the minor alleles within the MNP loci. MNP, multiple dispersed nucleotide polymorphism; BJ, Beijing; WH, Wuhan; HN, Hainan.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fmicb-16-1603255-g006.tif"/>
</fig>
</sec>
<sec id="S3.SS2.SSS4">
<title>3.2.4 Authentication of strains</title>
<p>The object of strain authentication is to determine whether the homonymous strains used at different periods (from the same laboratory or among different laboratories) are genetically consistent with their standard strain. The feasibility of MGV-Seq-based strain authentication was validated via the authentication of P6 and P8 homonymous lines. We took IRRI-P6 and IRRI-P8 as the standard strains, respectively. The homonymous lines of P6 included BJ-P6, HN-P6, WH-P6 and eight WH-P6 clones (CP6-1 &#x2013; CP6-8) and that of P8 included BJ-P8, HN-P8, WH-P8 and 14 WH-P8 clones (CP8-1 &#x2013; CP8-14). A pairwise comparison of the MNP major alleles of WH-P6, WH-P8 and their single-colony strains with IRRI-P6 and IRRI-P8, respectively, revealed that all the daughter lines showed 100% major allele similarity (MAS) with their standard strains IRRI-P6 and IRRI-P8, demonstrating that they were authenticated properly (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S13</xref>). Indeed, pathogenicity identification showed no significant differences among the three homogeneous WH-P6 clones (<xref ref-type="fig" rid="F5">Figure 5</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S10</xref>).</p>
<p>Overall, these results demonstrated that MGV-Seq can identify microorganisms at the strain level. Additionally, it can accurately detect low-abundance variations in cultured strains and uncultured host-contaminated samples (directly). Thus, MGV-Seq was readily available for daughter-line microevolution monitoring and strain authentication.</p>
</sec>
</sec>
</sec>
<sec id="S4" sec-type="discussion">
<title>4 Discussion</title>
<p>In this study, we focused on the development, performance evaluation, and demonstration of diverse applications of the MGV-Seq method. MGV-Seq genotyped a high-density panel of MNP markers within the genomes of target species using target enrichment and ultra-deep sequencing. Our results revealed 100% reproducibility and 100% accuracy for MGV-Seq in detecting major alleles and a comparable sensitivity but superior specificity in detecting genetic variations compared to the known exceptionally sensitive variant detection tool (<xref ref-type="bibr" rid="B42">Wilm et al., 2012</xref>; <xref ref-type="fig" rid="F2">Figure 2</xref>).</p>
<p>The prominent features of MGV-Seq can be attributed to several factors, as follows: First, MGV-Seq introduced as many markers as possible to minimize or avoid profiling errors; for example, MGV-Seq screened 213 MNP markers of the 4.49 Mbp genome of <italic>Xoo</italic>, forming a high-density panel of markers for a relatively small microbial genome (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S2</xref>). Second, in our study, MGV-Seq easily achieved ultra-high sequencing coverage&#x2014;8,425 folds per MNP marker (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S3</xref>). Such a dense panel of markers and such deep sequence coverage on the MNP markers would be ideal for genetic studies to accurately and sensitively detect and identify strains and analyze many low-frequency genetic variations with minimal profiling errors.</p>
<p>The prominent features of MGV-Seq make it markedly advantageous over existing methods for a wide range of applications, as illuminated in this study: (1) <italic>Strain identification</italic>: MGV-Seq achieved strain-level identification for pure culture and disentangled the genetic diversity of target microorganisms in uncultured diseased host samples, all with very high throughput and resolution. MGV-Seq overcomes the limitations of the electrophoresis band-based method commonly used for identifying cultured samples in terms of resolution, accuracy, and dependence on standard strains. It also shows advantages over the metagenomics-based method commonly used for identifying uncultured host samples regarding data volume, cost, analysis, and resolution. <bold>(2)</bold> <italic>Monitoring strain microevolution</italic>: MGV-Seq can accurately and sensitively capture low-frequency minor alleles that are difficult to capture using SNP genotyping-based methods (<xref ref-type="fig" rid="F3">Figure 3D</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S9</xref>). This characteristic confers MGV-Seq an advantage in detecting variation of strains at early periods, which is important for monitoring strain microevolution and new variants. However, it should be noted that phylogenetic discovery bias occurs when the molecular markers are derived from biased taxonomic sampling markers (<xref ref-type="bibr" rid="B29">Pearson et al., 2004</xref>). As shown in this study, the cluster dendrograms of IRRI-P1&#x2013;IRRI-P10 based on variations detected via MGV-Seq and WGS were completely consistent, suggesting that cluster analysis with multi-polymorphic markers is beneficial in avoiding the occurrence of this preference. (3) <italic>Monitoring strain purification</italic>: A pure culture is indispensable for infectious disease research. In this study, to demonstrate the feasibility of MGV-Seq-monitored strain purification, one round of single-colony isolation of WH-P6 and WH-P8, which MGV-Seq detected at 18 and 4 minor alleles (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S3</xref>), respectively, was performed using the spread plate and streak plate methods. All 22 clones showed 100% MAS with the parental strain, and the 15 clones harbored no minor alleles (<xref ref-type="supplementary-material" rid="TS1">Supplementary Table S3</xref>), all of which were homogeneous strains. This result showed that the clones were properly purified and that MGV-Seq could effectively monitor the strain purification process of heterogeneous samples. 4) Strain authentication. Heterogeneity of homonymous strains has been reported in several human pathogenic microbial species, such as <italic>Pseudomonas aeruginosa</italic> from 10 laboratories (<xref ref-type="bibr" rid="B5">Chandler et al., 2019</xref>) and the 13 daughter lines of the tuberculosis vaccine <italic>Mycobacterium bovis</italic> BCG (<xref ref-type="bibr" rid="B11">Garcia Pelayo et al., 2009</xref>). Our study reported genotype and phenotype discrepancies across the homonymous strains of <italic>X. oryzae</italic>. One solution to this serious problem is to introduce a mechanism of microorganism authentication, similar to that in human cell lines, to ensure that the same experimental materials are used in multiple experiments and across different laboratories. Therefore, the key index of technology for strain authentication was the high reproducibility and comparability of the genotyping results between experimental batches, ensuring result reproducibility and comparability across different laboratories. For the relatively small genome size of bacteria, compared with human cell lines, researchers may prefer the WGS method for strain authentication. However, in cases with large sample sizes, such as a study by (<xref ref-type="bibr" rid="B40">Song et al., 2023</xref>) which had more than 240 <italic>Xoo</italic> strains, or with large-genome microorganisms (fungi), the WGS method is still expensive (<xref ref-type="bibr" rid="B13">Guo et al., 2019</xref>). We have demonstrated that MGV-Seq showed 100% technical and batch reproducibility in detecting MNP major alleles (<xref ref-type="fig" rid="F2">Figure 2</xref>) and achieved a resolution approaching that of WGS-based methods. Using the barcodes and chips on existing platforms, for example, 768 barcodes and 550 Ion chips in the Ion Torrent sequencing platform, up to 1536 samples can be tested in a single sequencing run, enabling cost-effective MGV-Seq of large sample sets. The prominent features of MGV-Seq make it readily applicable to strain authentication (<xref ref-type="table" rid="T1">Table 1</xref>; <xref ref-type="supplementary-material" rid="TS1">Supplementary Table S13</xref>). Furthermore, because the heterogeneity of homonymous strains has also been observed in human pathogens, we strongly suggest MGV-Seq authentication of biological samples from other microorganisms.</p>
<table-wrap position="float" id="T1">
<label>TABLE 1</label>
<caption><p>Major allele similarity between single colony-derived strains and their parental strains.</p></caption>
<table cellspacing="5" cellpadding="5" frame="box" rules="all">
<thead>
<tr>
<td valign="top" align="left" style="color:#ffffff;background-color: #7f8080;">Single colony-derived strains</td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">Parental strains</td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">Number of compared loci</td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">Number of compared loci with same major alleles</td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">Major allele similarity</td>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">CP6-1</td>
<td valign="top" align="center">WH-P6</td>
<td valign="top" align="center">139</td>
<td valign="top" align="center">139</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP6-2</td>
<td valign="top" align="center">WH-P6</td>
<td valign="top" align="center">108</td>
<td valign="top" align="center">108</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP6-3</td>
<td valign="top" align="center">WH-P6</td>
<td valign="top" align="center">135</td>
<td valign="top" align="center">135</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP6-4</td>
<td valign="top" align="center">WH-P6</td>
<td valign="top" align="center">136</td>
<td valign="top" align="center">136</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP6-5</td>
<td valign="top" align="center">WH-P6</td>
<td valign="top" align="center">140</td>
<td valign="top" align="center">140</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP6-6</td>
<td valign="top" align="center">WH-P6</td>
<td valign="top" align="center">126</td>
<td valign="top" align="center">126</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP6-7</td>
<td valign="top" align="center">WH-P6</td>
<td valign="top" align="center">129</td>
<td valign="top" align="center">129</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP6-8</td>
<td valign="top" align="center">WH-P6</td>
<td valign="top" align="center">134</td>
<td valign="top" align="center">134</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-1</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">172</td>
<td valign="top" align="center">172</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-2</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">181</td>
<td valign="top" align="center">181</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-3</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">182</td>
<td valign="top" align="center">182</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-4</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">178</td>
<td valign="top" align="center">178</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-5</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">135</td>
<td valign="top" align="center">135</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-6</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">168</td>
<td valign="top" align="center">168</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-7</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">169</td>
<td valign="top" align="center">169</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-8</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">168</td>
<td valign="top" align="center">168</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-9</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">166</td>
<td valign="top" align="center">166</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-10</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">166</td>
<td valign="top" align="center">166</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-11</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">165</td>
<td valign="top" align="center">165</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-12</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">166</td>
<td valign="top" align="center">166</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-13</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">142</td>
<td valign="top" align="center">142</td>
<td valign="top" align="center">100%</td>
</tr>
<tr>
<td valign="top" align="left">CP8-14</td>
<td valign="top" align="center">WH-P8</td>
<td valign="top" align="center">154</td>
<td valign="top" align="center">154</td>
<td valign="top" align="center">100%</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn><p>CP6-1&#x223C;CP6-8 and CP8-1&#x223C;CP8-14: single colony-derived strains.</p></fn>
</table-wrap-foot>
</table-wrap>
<p>In addition to the applications that have been illuminated in this study, MGV-Seq is potentially applicable to high-throughput and low-cost detection of medical and forensic samples&#x2014;which typically have high host contamination or low biomass&#x2014;and for monitoring low-frequency genetic variations. Our study demonstrated that MGV-Seq could sensitively detect low-abundance variations directly from uncultured samples (<xref ref-type="fig" rid="F6">Figure 6</xref>). The targets of multiplex PCR can be in the thousands (<xref ref-type="bibr" rid="B20">Li et al., 2017</xref>). Thus, MGV-Seq is potentially feasible for detecting and typing multiple pathogen species in one panel in clinical samples, such as respiratory and intestinal tract samples, and plant quarantine samples, such as imported seeds and seedlings, by distributing the MNP loci to multiple pathogen species.</p>
<p>More importantly, the main steps for MGV-Seq are simple and compatible between sequencing platforms. Though MGV-Seq is highly dependent on the Ion Torrent technology in this study, MGV-Seq was readily available to any sequencing platform by adjusting the adaptors to the sequencing platforms on primers. In our previous study on plant variety identification, the Illumina high-throughput sequencing platform was used (<xref ref-type="bibr" rid="B7">Fang et al., 2021</xref>). Besides, we have already developed MNP panels for detecting SARS-CoV-2 based on Illumina and Beijing Genomics Institute sequencing platforms (data not shown). Sequencing data from any sequencing platform and laboratory can be shared through public databases. The genotypes can be derived and compared by our provided containerized tools, which can be compared freely and efficiently across laboratories.</p>
<p>Notably, several technical challenges should be noted. Firstly, it is well known that different sequencing platforms harbor their own sequencing bias (<xref ref-type="bibr" rid="B31">Quail et al., 2012</xref>, <xref ref-type="bibr" rid="B41">Suzuki et al., 2011</xref>), potentially leading to discordant genotypes for one sample. Therefore, it is necessary to consider comparing genotypes from different sequencing platforms. Furthermore, the ultra-sensitivity of MGV-Seq has made it sensitive to aerosol contamination and unavoidable technical problems of amplicon sequencing, such as PCR and sequencing errors, index hopping, and demultiplexing errors, which could lead to the generation of incorrect alleles. These alleles are often low-frequency and are confused with true low-frequency alleles. To identify the true low-frequency alleles, an optimal statistical analyzing pipeline (strand bias filter and statistical model) was adopted to remove incorrect alleles (see Materials and methods section). The high specificity of the minor alleles detected in the artificial DNA mixtures of the P6 and P8 strains proved the reliability of our statistical analysis pipeline for low-frequency minor allele detection.</p>
<p>In summary, the multifaceted capabilities of MGV-Seq underscore its potential to revolutionize microbial genomics across diverse fields. It integrates high - resolution strain - level discrimination, sensitivity to low - frequency variants, and cost - effectiveness. Its adaptability to multiple sequencing platforms enhance global standardization of strain authentication and comparative genomics. Beyond the demonstrated robust detection of pathogens in field samples, MGV-Seq&#x2019;s capacity to detect pathogen diversity in complex matrices (e.g., clinical samples with host contamination, low-biomass forensic specimens) positions it as a versatile tool for precision medicine, public health surveillance, and biosecurity. Crucially, MGV-Seq&#x2019;s emphasis on reproducibility and cross-laboratory comparability highlights the need for standardized microbial authentication, akin to human cell line authentication systems. Its framework for tracking genetic variation at early stages of evolution could also inform strategies against antimicrobial resistance or pathogen adaptation. In the future, as a foundational technology for microbiology, MGV-Seq is poised to fostering global collaboration through shared databases and containerized tools, seamless adaptation to new organisms and emerging challenges.</p>
</sec>
</body>
<back>
<sec id="S5" sec-type="data-availability">
<title>Data availability statement</title>
<p>The data supporting this study&#x2019;s findings are openly available in the NCBI BioProject database (<ext-link ext-link-type="uri" xlink:href="https://www.ncbi.nlm.nih.gov/bioproject/">https://www.ncbi.nlm.nih.gov/bioproject/</ext-link>) at accession number <ext-link ext-link-type="DDBJ/EMBL/GenBank" xlink:href="PRJNA679673">PRJNA679673</ext-link>. In addition, the codes generated during this study are available at: <ext-link ext-link-type="uri" xlink:href="https://github.com/SystemsBiologyOfJianghanUniversity/MGV-Seq">https://github.com/SystemsBiologyOfJianghanUniversity/MGV-Seq</ext-link>.</p>
</sec>
<sec id="S6" sec-type="author-contributions">
<title>Author contributions</title>
<p>LL: Software, Writing &#x2013; original draft, Formal Analysis, Methodology. WK: Investigation, Methodology, Writing &#x2013; original draft, Data curation, Validation, Visualization. JS: Data curation, Writing &#x2013; original draft, Validation, Visualization. YJ: Supervision, Writing &#x2013; review &#x0026; editing, Resources. TL: Formal Analysis, Writing &#x2013; original draft, Validation. ZX: Resources, Writing &#x2013; original draft. JZ: Formal Analysis, Writing &#x2013; original draft, Data curation, Validation. ZF: Formal Analysis, Writing &#x2013; original draft. LC: Formal Analysis, Writing &#x2013; original draft, Validation. SF: Investigation, Methodology, Writing &#x2013; original draft. HS: Investigation, Methodology, Writing &#x2013; original draft. HX: Data curation, Writing &#x2013; original draft. BZ: Investigation, Methodology, Writing &#x2013; original draft. BF: Project administration, Writing &#x2013; original draft. HP: Conceptualization, Writing &#x2013; review &#x0026; editing, Project administration, Resources, Supervision. LG: Conceptualization, Funding acquisition, Writing &#x2013; review &#x0026; editing, Project administration, Supervision.</p>
</sec>
<sec id="S7" sec-type="funding-information">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research and/or publication of this article. This work was financially supported by the Internal Projects of Jianghan University (grant no. 2023KJZX44).</p>
</sec>
<ack><p>We thank Oliva Ricardo for providing the DNA samples of IRRI-P1&#x2013;IRRI-P10, Professor Wenxue Zhai for providing the BJ-P1&#x2013;BJ-P10 strains and Zhihun Xia for providing the HN-P1&#x2013;HN-P10 strains.</p>
</ack>
<sec id="S8" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>HP was employed by the Mingliao Biotechnology Co., Ltd. LG was employed by the Wuhan Zhongwei Gene Technology Co., Ltd. The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="S9" sec-type="ai-statement">
<title>Generative AI statement</title>
<p>The authors declare that no Generative AI was used in the creation of this manuscript.</p>
</sec>
<sec id="S10" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="S11" sec-type="supplementary-material">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fmicb.2025.1603255/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fmicb.2025.1603255/full#supplementary-material</ext-link></p>
<supplementary-material xlink:href="Table_1.xlsx" id="TS1" mimetype="application/vnd.openxmlformats-officedocument.spreadsheetml.sheet" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image_1.TIF" id="FS1" mimetype="image/tiff" xmlns:xlink="http://www.w3.org/1999/xlink">
<label>Supplementary Figure S1</label>
<caption><p>The transmission time and route of the <italic>Xanthomonas oryzae</italic> pv. <italic>oryzae</italic> strains used in this study. The strains enclosed in the black box were used in this study.</p></caption>
</supplementary-material>
<supplementary-material xlink:href="Image_2.TIF" id="FS2" mimetype="image/tiff" xmlns:xlink="http://www.w3.org/1999/xlink">
<label>Supplementary Figure S2</label>
<caption><p>The customized computational pipeline for MNP genotyping includes calling candidate alleles, identifying major alleles, and detecting true minor alleles. MNP, multiple dispersed nucleotide polymorphism; SNPs, single nucleotide polymorphisms.</p></caption>
</supplementary-material>
<supplementary-material xlink:href="Image_3.TIF" id="FS3" mimetype="image/tiff" xmlns:xlink="http://www.w3.org/1999/xlink">
<label>Supplementary Figure S3</label>
<caption><p>The heatmap of pairwise comparisons of homonymous strains.</p></caption>
</supplementary-material>
</sec>
<fn-group>
<fn id="footnote1">
<label>1</label>
<p><ext-link ext-link-type="uri" xlink:href="https://ampliseq.com/">https://ampliseq.com/</ext-link></p></fn>
<fn id="footnote2">
<label>2</label>
<p><ext-link ext-link-type="uri" xlink:href="http://hannonlab.cshl.edu/fastx_toolkit/index.html">http://hannonlab.cshl.edu/fastx_toolkit/index.html</ext-link></p></fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Andreu</surname> <given-names>N.</given-names></name> <name><surname>Gibert</surname> <given-names>I.</given-names></name></person-group> (<year>2008</year>). <article-title>Cell population heterogeneity in Mycobacterium tuberculosis H37Rv.</article-title> <source><italic>Tuberculosis</italic></source> <volume>88</volume> <fpage>553</fpage>&#x2013;<lpage>559</lpage>. <pub-id pub-id-type="doi">10.1016/j.tube.2008.03.005</pub-id> <pub-id pub-id-type="pmid">18502178</pub-id></citation></ref>
<ref id="B2"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Antil</surname> <given-names>S.</given-names></name> <name><surname>Abraham</surname> <given-names>J. S.</given-names></name> <name><surname>Sripoorna</surname> <given-names>S.</given-names></name> <name><surname>Maurya</surname> <given-names>S.</given-names></name> <name><surname>Dagar</surname> <given-names>J.</given-names></name> <name><surname>Makhija</surname> <given-names>S.</given-names></name><etal/></person-group> (<year>2023</year>). <article-title>DNA barcoding, an effective tool for species identification: A review.</article-title> <source><italic>Mol. Biol. Rep.</italic></source> <volume>50</volume> <fpage>761</fpage>&#x2013;<lpage>775</lpage>. <pub-id pub-id-type="doi">10.1007/s11033-022-08015-7</pub-id> <pub-id pub-id-type="pmid">36308581</pub-id></citation></ref>
<ref id="B3"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bayer</surname> <given-names>M. M.</given-names></name> <name><surname>Rapazote-Flores</surname> <given-names>P.</given-names></name> <name><surname>Ganal</surname> <given-names>M.</given-names></name> <name><surname>Hedley</surname> <given-names>P. E.</given-names></name> <name><surname>Macaulay</surname> <given-names>M.</given-names></name> <name><surname>Plieske</surname> <given-names>J.</given-names></name><etal/></person-group> (<year>2017</year>). <article-title>Development and evaluation of a barley 50k iSelect SNP array.</article-title> <source><italic>Front. Plant Sci.</italic></source> <volume>8</volume>:<fpage>1792</fpage>. <pub-id pub-id-type="doi">10.3389/fpls.2017.01792</pub-id> <pub-id pub-id-type="pmid">29089957</pub-id></citation></ref>
<ref id="B4"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Catara</surname> <given-names>V.</given-names></name> <name><surname>Cubero</surname> <given-names>J.</given-names></name> <name><surname>Pothier</surname> <given-names>J. F.</given-names></name> <name><surname>Bosis</surname> <given-names>E.</given-names></name> <name><surname>Bragard</surname> <given-names>C.</given-names></name> <name><surname>&#x0110;ermi&#x0107;</surname> <given-names>E.</given-names></name><etal/></person-group> (<year>2021</year>). <article-title>Trends in molecular diagnosis and diversity studies for phytosanitary regulated xanthomonas.</article-title> <source><italic>Microorganisms</italic></source> <volume>9</volume>:<fpage>862</fpage>. <pub-id pub-id-type="doi">10.3390/microorganisms9040862</pub-id> <pub-id pub-id-type="pmid">33923763</pub-id></citation></ref>
<ref id="B5"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chandler</surname> <given-names>C. E.</given-names></name> <name><surname>Horspool</surname> <given-names>A. M.</given-names></name> <name><surname>Hill</surname> <given-names>P. J.</given-names></name> <name><surname>Wozniak</surname> <given-names>D. J.</given-names></name> <name><surname>Schertzer</surname> <given-names>J. W.</given-names></name> <name><surname>Rasko</surname> <given-names>D. A.</given-names></name><etal/></person-group> (<year>2019</year>). <article-title>Genomic and phenotypic diversity among ten laboratory isolates of <italic>Pseudomonas aeruginosa</italic> PAO1.</article-title> <source><italic>J. Bacteriol.</italic></source> <volume>201</volume>:<fpage>e00595-18</fpage>. <pub-id pub-id-type="doi">10.1128/JB.00595-18</pub-id> <pub-id pub-id-type="pmid">30530517</pub-id></citation></ref>
<ref id="B6"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fang</surname> <given-names>L.</given-names></name> <name><surname>Wang</surname> <given-names>Q.</given-names></name> <name><surname>Hu</surname> <given-names>Y.</given-names></name> <name><surname>Jia</surname> <given-names>Y.</given-names></name> <name><surname>Chen</surname> <given-names>J.</given-names></name> <name><surname>Liu</surname> <given-names>B.</given-names></name><etal/></person-group> (<year>2017</year>). <article-title>Genomic analyses in cotton identify signatures of selection and loci associated with fiber quality and yield traits.</article-title> <source><italic>Nat. Genet.</italic></source> <volume>49</volume> <fpage>1089</fpage>&#x2013;<lpage>1098</lpage>. <pub-id pub-id-type="doi">10.1038/ng.3887</pub-id> <pub-id pub-id-type="pmid">28581501</pub-id></citation></ref>
<ref id="B7"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fang</surname> <given-names>Z.</given-names></name> <name><surname>Li</surname> <given-names>L.</given-names></name> <name><surname>Zhou</surname> <given-names>J.</given-names></name> <name><surname>You</surname> <given-names>A.</given-names></name> <name><surname>Gao</surname> <given-names>L.</given-names></name> <name><surname>Li</surname> <given-names>T.</given-names></name><etal/></person-group> (<year>2021</year>). <article-title>Multiple nucleotide polymorphism DNA markers for the accurate evaluation of genetic variations.</article-title> <source><italic>bioRxiv [Prperint]</italic></source> <pub-id pub-id-type="doi">10.1101/2021.03.09.434561</pub-id></citation></ref>
<ref id="B8"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Feng</surname> <given-names>C.</given-names></name> <name><surname>Xu</surname> <given-names>F.</given-names></name> <name><surname>Li</surname> <given-names>L.</given-names></name> <name><surname>Zhang</surname> <given-names>J.</given-names></name> <name><surname>Wang</surname> <given-names>J.</given-names></name> <name><surname>Li</surname> <given-names>Y.</given-names></name><etal/></person-group> (<year>2023</year>). <article-title>Biological control of <italic>Fusarium</italic> crown rot of wheat with <italic>Chaetomium globosum</italic> 12XP1-2-3 and its effects on rhizosphere microorganisms.</article-title> <source><italic>Front. Microbiol.</italic></source> <volume>14</volume>:<fpage>1133025</fpage>. <pub-id pub-id-type="doi">10.3389/fmicb.2023.1133025</pub-id> <pub-id pub-id-type="pmid">37077244</pub-id></citation></ref>
<ref id="B9"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Feng</surname> <given-names>W.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Huang</surname> <given-names>L.</given-names></name> <name><surname>Feng</surname> <given-names>C.</given-names></name> <name><surname>Chu</surname> <given-names>Z.</given-names></name> <name><surname>Ding</surname> <given-names>X.</given-names></name><etal/></person-group> (<year>2015</year>). <article-title>Genomic-associated markers and comparative genome maps of <italic>Xanthomonas oryzae</italic> pv. oryzae and <italic>X. oryzae pv</italic>. oryzicola.</article-title> <source><italic>World J. Microbiol. Biotechnol.</italic></source> <volume>31</volume> <fpage>1353</fpage>&#x2013;<lpage>1359</lpage>. <pub-id pub-id-type="doi">10.1007/s11274-015-1883-5</pub-id> <pub-id pub-id-type="pmid">26093644</pub-id></citation></ref>
<ref id="B10"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gao</surname> <given-names>L.</given-names></name> <name><surname>Li</surname> <given-names>L.</given-names></name> <name><surname>Fang</surname> <given-names>B.</given-names></name> <name><surname>Fang</surname> <given-names>Z.</given-names></name> <name><surname>Xiang</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>M.</given-names></name><etal/></person-group> (<year>2023</year>). <article-title>Carryover contamination-controlled amplicon sequencing workflow for accurate qualitative and quantitative detection of pathogens: A case study on SARS-CoV-2.</article-title> <source><italic>Microbiol. Spectr.</italic></source> <volume>11</volume>:<fpage>e0020623</fpage>. <pub-id pub-id-type="doi">10.1128/spectrum.00206-23</pub-id> <pub-id pub-id-type="pmid">37098913</pub-id></citation></ref>
<ref id="B11"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Garcia Pelayo</surname> <given-names>M. C.</given-names></name> <name><surname>Uplekar</surname> <given-names>S.</given-names></name> <name><surname>Keniry</surname> <given-names>A.</given-names></name> <name><surname>Mendoza Lopez</surname> <given-names>P.</given-names></name> <name><surname>Garnier</surname> <given-names>T.</given-names></name> <name><surname>Nunez Garcia</surname> <given-names>J.</given-names></name><etal/></person-group> (<year>2009</year>). <article-title>A comprehensive survey of single nucleotide polymorphisms (SNPs) across <italic>Mycobacterium bovis</italic> strains and <italic>M. bovis</italic> BCG vaccine strains refines the genealogy and defines a minimal set of SNPs that separate virulent <italic>M. bovis</italic> strains and <italic>M. bovis</italic> BCG strains.</article-title> <source><italic>Infect. Immun.</italic></source> <volume>77</volume> <fpage>2230</fpage>&#x2013;<lpage>2238</lpage>. <pub-id pub-id-type="doi">10.1128/IAI.01099-08</pub-id> <pub-id pub-id-type="pmid">19289514</pub-id></citation></ref>
<ref id="B12"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Grissa</surname> <given-names>I.</given-names></name> <name><surname>Bouchon</surname> <given-names>P.</given-names></name> <name><surname>Pourcel</surname> <given-names>C.</given-names></name> <name><surname>Vergnaud</surname> <given-names>G.</given-names></name></person-group> (<year>2008</year>). <article-title>On-line resources for bacterial micro-evolution studies using MLVA or CRISPR typing.</article-title> <source><italic>Biochimie</italic></source> <volume>90</volume> <fpage>660</fpage>&#x2013;<lpage>668</lpage>. <pub-id pub-id-type="doi">10.1016/j.biochi.2007.07.014</pub-id> <pub-id pub-id-type="pmid">17822824</pub-id></citation></ref>
<ref id="B13"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>Z. F.</given-names></name> <name><surname>Wang</surname> <given-names>H. W.</given-names></name> <name><surname>Tao</surname> <given-names>J. J.</given-names></name> <name><surname>Ren</surname> <given-names>Y. H.</given-names></name> <name><surname>Xu</surname> <given-names>C.</given-names></name> <name><surname>Wu</surname> <given-names>K. S.</given-names></name><etal/></person-group> (<year>2019</year>). <article-title>Development of multiple SNP marker panels affordable to breeders through genotyping by target sequencing (GBTS) in maize.</article-title> <source><italic>Mol. Breed.</italic></source> <volume>39</volume>:<fpage>37</fpage>. <pub-id pub-id-type="doi">10.1007/s11032-019-0940-4</pub-id></citation></ref>
<ref id="B14"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hajri</surname> <given-names>A.</given-names></name> <name><surname>Brin</surname> <given-names>C.</given-names></name> <name><surname>Zhao</surname> <given-names>S.</given-names></name> <name><surname>David</surname> <given-names>P.</given-names></name> <name><surname>Feng</surname> <given-names>J. X.</given-names></name> <name><surname>Koebnik</surname> <given-names>R.</given-names></name><etal/></person-group> (<year>2012</year>). <article-title>Multilocus sequence analysis and type III effector repertoire mining provide new insights into the evolutionary history and virulence of <italic>Xanthomonas oryzae</italic>.</article-title> <source><italic>Mol. Plant Pathol.</italic></source> <volume>13</volume> <fpage>288</fpage>&#x2013;<lpage>302</lpage>. <pub-id pub-id-type="doi">10.1111/j.1364-3703.2011.00745.x</pub-id> <pub-id pub-id-type="pmid">21929565</pub-id></citation></ref>
<ref id="B15"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Huang</surname> <given-names>Y.</given-names></name> <name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Zheng</surname> <given-names>C.</given-names></name> <name><surname>Shen</surname> <given-names>C.</given-names></name></person-group> (<year>2017</year>). <article-title>Investigation of cross-contamination and misidentification of 278 widely used tumor cell lines.</article-title> <source><italic>PLoS One</italic></source> <volume>12</volume>:<fpage>e0170384</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0170384</pub-id> <pub-id pub-id-type="pmid">28107433</pub-id></citation></ref>
<ref id="B16"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jongman</surname> <given-names>M.</given-names></name> <name><surname>Carmichael</surname> <given-names>P. C.</given-names></name> <name><surname>Bill</surname> <given-names>M.</given-names></name></person-group> (<year>2020</year>). <article-title>Technological advances in phytopathogen detection and metagenome profiling techniques.</article-title> <source><italic>Curr. Microbiol.</italic></source> <volume>77</volume> <fpage>675</fpage>&#x2013;<lpage>681</lpage>. <pub-id pub-id-type="doi">10.1007/s00284-020-01881-z</pub-id> <pub-id pub-id-type="pmid">31960092</pub-id></citation></ref>
<ref id="B17"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Langmead</surname> <given-names>B.</given-names></name> <name><surname>Salzberg</surname> <given-names>S. L.</given-names></name></person-group> (<year>2012</year>). <article-title>Fast gapped-read alignment with Bowtie 2.</article-title> <source><italic>Nat. Methods</italic></source> <volume>9</volume> <fpage>357</fpage>&#x2013;<lpage>359</lpage>. <pub-id pub-id-type="doi">10.1038/nmeth.1923</pub-id> <pub-id pub-id-type="pmid">22388286</pub-id></citation></ref>
<ref id="B18"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lee</surname> <given-names>B. M.</given-names></name> <name><surname>Park</surname> <given-names>Y. J.</given-names></name> <name><surname>Park</surname> <given-names>D. S.</given-names></name> <name><surname>Kang</surname> <given-names>H. W.</given-names></name> <name><surname>Kim</surname> <given-names>J. G.</given-names></name> <name><surname>Song</surname> <given-names>E. S.</given-names></name><etal/></person-group> (<year>2005</year>). <article-title>The genome sequence of <italic>Xanthomonas oryzae</italic> pathovar oryzae KACC10331, the bacterial blight pathogen of rice.</article-title> <source><italic>Nucleic Acids Res.</italic></source> <volume>33</volume> <fpage>577</fpage>&#x2013;<lpage>586</lpage>. <pub-id pub-id-type="doi">10.1093/nar/gki206</pub-id> <pub-id pub-id-type="pmid">15673718</pub-id></citation></ref>
<ref id="B19"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>G.</given-names></name></person-group> (<year>2016</year>). <article-title>A new model calling procedure for Illumina BeadArray data.</article-title> <source><italic>BMC Genet.</italic></source> <volume>17</volume>:<fpage>90</fpage>. <pub-id pub-id-type="doi">10.1186/s12863-016-0398-x</pub-id> <pub-id pub-id-type="pmid">27343118</pub-id></citation></ref>
<ref id="B20"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>L.</given-names></name> <name><surname>Fang</surname> <given-names>Z.</given-names></name> <name><surname>Zhou</surname> <given-names>J.</given-names></name> <name><surname>Chen</surname> <given-names>H.</given-names></name> <name><surname>Hu</surname> <given-names>Z.</given-names></name> <name><surname>Gao</surname> <given-names>L.</given-names></name><etal/></person-group> (<year>2017</year>). <article-title>An accurate and efficient method for large-scale SSR genotyping and applications.</article-title> <source><italic>Nucleic Acids Res.</italic></source> <volume>45</volume>:<fpage>e88</fpage>. <pub-id pub-id-type="doi">10.1093/nar/gkx093</pub-id> <pub-id pub-id-type="pmid">28184437</pub-id></citation></ref>
<ref id="B21"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Mi</surname> <given-names>Y.</given-names></name> <name><surname>Mueller</surname> <given-names>T.</given-names></name> <name><surname>Kreibich</surname> <given-names>S.</given-names></name> <name><surname>Williams</surname> <given-names>E. G.</given-names></name> <name><surname>Van Drogen</surname> <given-names>A.</given-names></name><etal/></person-group> (<year>2019</year>). <article-title>Multi-omic measurements of heterogeneity in HeLa cells across laboratories.</article-title> <source><italic>Nat. Biotechnol.</italic></source> <volume>37</volume> <fpage>314</fpage>&#x2013;<lpage>322</lpage>. <pub-id pub-id-type="doi">10.1038/s41587-019-0037-y</pub-id> <pub-id pub-id-type="pmid">30778230</pub-id></citation></ref>
<ref id="B22"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lorsch</surname> <given-names>J. R.</given-names></name> <name><surname>Collins</surname> <given-names>F. S.</given-names></name> <name><surname>Lippincott-Schwartz</surname> <given-names>J.</given-names></name></person-group> (<year>2014</year>). <article-title>Cell biology. fixing problems with cell lines.</article-title> <source><italic>Science</italic></source> <volume>346</volume> <fpage>1452</fpage>&#x2013;<lpage>1453</lpage>. <pub-id pub-id-type="doi">10.1126/science.1259110</pub-id> <pub-id pub-id-type="pmid">25525228</pub-id></citation></ref>
<ref id="B23"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Machado</surname> <given-names>G. E.</given-names></name> <name><surname>Matsumoto</surname> <given-names>C. K.</given-names></name> <name><surname>Chimara</surname> <given-names>E.</given-names></name> <name><surname>Duarte Rda</surname> <given-names>S.</given-names></name> <name><surname>de Freitas</surname> <given-names>D.</given-names></name> <name><surname>Palaci</surname> <given-names>M.</given-names></name><etal/></person-group> (<year>2014</year>). <article-title>Multilocus sequence typing scheme versus pulsed-field gel electrophoresis for typing Mycobacterium abscessus isolates.</article-title> <source><italic>J. Clin. Microbiol.</italic></source> <volume>52</volume> <fpage>2881</fpage>&#x2013;<lpage>2891</lpage>. <pub-id pub-id-type="doi">10.1128/JCM.00688-14</pub-id> <pub-id pub-id-type="pmid">24899019</pub-id></citation></ref>
<ref id="B24"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Maiden</surname> <given-names>M. C.</given-names></name></person-group> (<year>2006</year>). <article-title>Multilocus sequence typing of bacteria.</article-title> <source><italic>Annu. Rev. Microbiol.</italic></source> <volume>60</volume> <fpage>561</fpage>&#x2013;<lpage>588</lpage>. <pub-id pub-id-type="doi">10.1146/annurev.micro.59.030804.121325</pub-id> <pub-id pub-id-type="pmid">16774461</pub-id></citation></ref>
<ref id="B25"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mansfield</surname> <given-names>J.</given-names></name> <name><surname>Genin</surname> <given-names>S.</given-names></name> <name><surname>Magori</surname> <given-names>S.</given-names></name> <name><surname>Citovsky</surname> <given-names>V.</given-names></name> <name><surname>Sriariyanum</surname> <given-names>M.</given-names></name> <name><surname>Ronald</surname> <given-names>P.</given-names></name><etal/></person-group> (<year>2012</year>). <article-title>Top 10 plant pathogenic bacteria in molecular plant pathology.</article-title> <source><italic>Mol. Plant Pathol.</italic></source> <volume>13</volume> <fpage>614</fpage>&#x2013;<lpage>629</lpage>. <pub-id pub-id-type="doi">10.1111/j.1364-3703.2012.00804.x</pub-id> <pub-id pub-id-type="pmid">22672649</pub-id></citation></ref>
<ref id="B26"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Masters</surname> <given-names>J. R.</given-names></name></person-group> (<year>2012</year>). <article-title>Cell-line authentication: End the scandal of false cell lines.</article-title> <source><italic>Nature</italic></source> <volume>492</volume>:<fpage>186</fpage>. <pub-id pub-id-type="doi">10.1038/492186a</pub-id> <pub-id pub-id-type="pmid">23235867</pub-id></citation></ref>
<ref id="B27"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Newberry</surname> <given-names>E.</given-names></name> <name><surname>Bhandari</surname> <given-names>R.</given-names></name> <name><surname>Kemble</surname> <given-names>J.</given-names></name> <name><surname>Sikora</surname> <given-names>E.</given-names></name> <name><surname>Potnis</surname> <given-names>N.</given-names></name></person-group> (<year>2020</year>). <article-title>Genome-resolved metagenomics to study co-occurrence patterns and intraspecific heterogeneity among plant pathogen metapopulations.</article-title> <source><italic>Environ. Microbiol.</italic></source> <volume>22</volume> <fpage>2693</fpage>&#x2013;<lpage>2708</lpage>. <pub-id pub-id-type="doi">10.1111/1462-2920.14989</pub-id> <pub-id pub-id-type="pmid">32207218</pub-id></citation></ref>
<ref id="B28"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ni&#x00F1;o-Liu</surname> <given-names>D. O.</given-names></name> <name><surname>Ronald</surname> <given-names>P. C.</given-names></name> <name><surname>Bogdanove</surname> <given-names>A. J.</given-names></name></person-group> (<year>2006</year>). <article-title><italic>Xanthomonas oryzae</italic> pathovars: Model pathogens of a model crop.</article-title> <source><italic>Mol. Plant Pathol.</italic></source> <volume>7</volume> <fpage>303</fpage>&#x2013;<lpage>324</lpage>. <pub-id pub-id-type="doi">10.1111/j.1364-3703.2006.00344.x</pub-id> <pub-id pub-id-type="pmid">20507449</pub-id></citation></ref>
<ref id="B29"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pearson</surname> <given-names>T.</given-names></name> <name><surname>Busch</surname> <given-names>J. D.</given-names></name> <name><surname>Ravel</surname> <given-names>J.</given-names></name> <name><surname>Read</surname> <given-names>T. D.</given-names></name> <name><surname>Rhoton</surname> <given-names>S. D.</given-names></name> <name><surname>U&#x2019;Ren</surname> <given-names>J. M.</given-names></name><etal/></person-group> (<year>2004</year>). <article-title>Phylogenetic discovery bias in <italic>Bacillus anthracis</italic> using single-nucleotide polymorphisms from whole-genome sequencing.</article-title> <source><italic>Proc. Natl. Acad. Sci. U. S. A.</italic></source> <volume>101</volume> <fpage>13536</fpage>&#x2013;<lpage>13541</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.0403844101</pub-id> <pub-id pub-id-type="pmid">15347815</pub-id></citation></ref>
<ref id="B30"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Poulin</surname> <given-names>L.</given-names></name> <name><surname>Grygiel</surname> <given-names>P.</given-names></name> <name><surname>Magne</surname> <given-names>M.</given-names></name> <name><surname>Gagnevin</surname> <given-names>L.</given-names></name> <name><surname>Rodriguez-R</surname> <given-names>L. M.</given-names></name> <name><surname>Forero Serna</surname> <given-names>N.</given-names></name><etal/></person-group> (<year>2015</year>). <article-title>New multilocus variable-number tandem-repeat analysis tool for surveillance and local epidemiology of bacterial leaf blight and bacterial leaf streak of rice caused by <italic>Xanthomonas oryzae</italic>.</article-title> <source><italic>Appl. Environ. Microbiol.</italic></source> <volume>81</volume> <fpage>688</fpage>&#x2013;<lpage>698</lpage>. <pub-id pub-id-type="doi">10.1128/AEM.02768-14</pub-id> <pub-id pub-id-type="pmid">25398857</pub-id></citation></ref>
<ref id="B31"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Quail</surname> <given-names>M. A.</given-names></name> <name><surname>Smith</surname> <given-names>M.</given-names></name> <name><surname>Coupland</surname> <given-names>P.</given-names></name> <name><surname>Otto</surname> <given-names>T. D.</given-names></name> <name><surname>Harris</surname> <given-names>S. R.</given-names></name> <name><surname>Connor</surname> <given-names>T. R.</given-names></name><etal/></person-group> (<year>2012</year>). <article-title>A tale of three next generation sequencing platforms: Comparison of ion torrent, pacific biosciences and Illumina MiSeq sequencers.</article-title> <source><italic>BMC Genom.</italic></source> <volume>13</volume>:<fpage>341</fpage>. <pub-id pub-id-type="doi">10.1186/1471-2164-13-341</pub-id> <pub-id pub-id-type="pmid">22827831</pub-id></citation></ref>
<ref id="B32"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Quibod</surname> <given-names>I. L.</given-names></name> <name><surname>Atieza-Grande</surname> <given-names>G.</given-names></name> <name><surname>Oreiro</surname> <given-names>E. G.</given-names></name> <name><surname>Palmos</surname> <given-names>D.</given-names></name> <name><surname>Nguyen</surname> <given-names>M. H.</given-names></name> <name><surname>Coronejo</surname> <given-names>S. T.</given-names></name><etal/></person-group> (<year>2020</year>). <article-title>The green revolution shaped the population structure of the rice pathogen <italic>Xanthomonas oryzae</italic> pv. oryzae.</article-title> <source><italic>ISME J.</italic></source> <volume>14</volume> <fpage>492</fpage>&#x2013;<lpage>505</lpage>. <pub-id pub-id-type="doi">10.1038/s41396-019-0545-2</pub-id> <pub-id pub-id-type="pmid">31666657</pub-id></citation></ref>
<ref id="B33"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Quibod</surname> <given-names>I. L.</given-names></name> <name><surname>Perez-Quintero</surname> <given-names>A.</given-names></name> <name><surname>Booher</surname> <given-names>N. J.</given-names></name> <name><surname>Dossa</surname> <given-names>G. S.</given-names></name> <name><surname>Grande</surname> <given-names>G.</given-names></name> <name><surname>Szurek</surname> <given-names>B.</given-names></name><etal/></person-group> (<year>2016</year>). <article-title>Effector diversification contributes to <italic>Xanthomonas oryzae</italic> pv. oryzae phenotypic adaptation in a semi-isolated environment.</article-title> <source><italic>Sci. Rep.</italic></source> <volume>6</volume>:<fpage>34137</fpage>. <pub-id pub-id-type="doi">10.1038/srep34137</pub-id> <pub-id pub-id-type="pmid">27667260</pub-id></citation></ref>
<ref id="B34"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rademaker</surname> <given-names>J. L.</given-names></name> <name><surname>Louws</surname> <given-names>F. J.</given-names></name> <name><surname>Schultz</surname> <given-names>M. H.</given-names></name> <name><surname>Rossbach</surname> <given-names>U.</given-names></name> <name><surname>Vauterin</surname> <given-names>L.</given-names></name> <name><surname>Swings</surname> <given-names>J.</given-names></name><etal/></person-group> (<year>2005</year>). <article-title>A comprehensive species to strain taxonomic framework for xanthomonas.</article-title> <source><italic>Phytopathology</italic></source> <volume>95</volume> <fpage>1098</fpage>&#x2013;<lpage>1111</lpage>. <pub-id pub-id-type="doi">10.1094/PHYTO-95-1098</pub-id> <pub-id pub-id-type="pmid">18943308</pub-id></citation></ref>
<ref id="B35"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Roychowdhury</surname> <given-names>T.</given-names></name> <name><surname>Singh</surname> <given-names>V. K.</given-names></name> <name><surname>Bhattacharya</surname> <given-names>A.</given-names></name></person-group> (<year>2019</year>). <article-title>Classification of pathogenic microbes using a minimal set of single nucleotide polymorphisms derived from whole genome sequences.</article-title> <source><italic>Genomics</italic></source> <volume>111</volume> <fpage>205</fpage>&#x2013;<lpage>211</lpage>. <pub-id pub-id-type="doi">10.1016/j.ygeno.2018.02.004</pub-id> <pub-id pub-id-type="pmid">29432978</pub-id></citation></ref>
<ref id="B36"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Saltykova</surname> <given-names>A.</given-names></name> <name><surname>Wuyts</surname> <given-names>V.</given-names></name> <name><surname>Mattheus</surname> <given-names>W.</given-names></name> <name><surname>Bertrand</surname> <given-names>S.</given-names></name> <name><surname>Roosens</surname> <given-names>N. H. C.</given-names></name> <name><surname>Marchal</surname> <given-names>K.</given-names></name><etal/></person-group> (<year>2018</year>). <article-title>Comparison of SNP-based subtyping workflows for bacterial isolates using WGS data, applied to <italic>Salmonella enterica</italic> serotype <italic>Typhimurium</italic> and serotype 1,4,[5],12:i.</article-title> <source><italic>PLoS One</italic></source> <volume>13</volume>:<fpage>e0192504</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0192504</pub-id> <pub-id pub-id-type="pmid">29408896</pub-id></citation></ref>
<ref id="B37"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Salzberg</surname> <given-names>S. L.</given-names></name> <name><surname>Sommer</surname> <given-names>D. D.</given-names></name> <name><surname>Schatz</surname> <given-names>M. C.</given-names></name> <name><surname>Phillippy</surname> <given-names>A. M.</given-names></name> <name><surname>Rabinowicz</surname> <given-names>P. D.</given-names></name> <name><surname>Tsuge</surname> <given-names>S.</given-names></name><etal/></person-group> (<year>2008</year>). <article-title>Genome sequence and rapid evolution of the rice pathogen <italic>Xanthomonas oryzae</italic> pv. oryzae PXO99A.</article-title> <source><italic>BMC Genom.</italic></source> <volume>9</volume>:<fpage>204</fpage>. <pub-id pub-id-type="doi">10.1186/1471-2164-9-204</pub-id> <pub-id pub-id-type="pmid">18452608</pub-id></citation></ref>
<ref id="B38"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sasaki</surname> <given-names>S.</given-names></name> <name><surname>Yoshinari</surname> <given-names>K.</given-names></name> <name><surname>Uchiyama</surname> <given-names>K.</given-names></name> <name><surname>Takeda</surname> <given-names>M.</given-names></name> <name><surname>Kojima</surname> <given-names>T.</given-names></name></person-group> (<year>2018</year>). <article-title>Relationship between call rate per individual and genotyping accuracy of bovine single-nucleotide polymorphism array using deoxyribonucleic acid of various qualities.</article-title> <source><italic>Anim. Sci. J.</italic></source> <volume>89</volume> <fpage>1533</fpage>&#x2013;<lpage>1539</lpage>. <pub-id pub-id-type="doi">10.1111/asj.13110</pub-id> <pub-id pub-id-type="pmid">30230122</pub-id></citation></ref>
<ref id="B39"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Solden</surname> <given-names>L.</given-names></name> <name><surname>Lloyd</surname> <given-names>K.</given-names></name> <name><surname>Wrighton</surname> <given-names>K.</given-names></name></person-group> (<year>2016</year>). <article-title>The bright side of microbial dark matter: Lessons learned from the uncultivated majority.</article-title> <source><italic>Curr. Opin. Microbiol.</italic></source> <volume>31</volume> <fpage>217</fpage>&#x2013;<lpage>226</lpage>. <pub-id pub-id-type="doi">10.1016/j.mib.2016.04.020</pub-id> <pub-id pub-id-type="pmid">27196505</pub-id></citation></ref>
<ref id="B40"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Song</surname> <given-names>Z.</given-names></name> <name><surname>Zheng</surname> <given-names>J.</given-names></name> <name><surname>Zhao</surname> <given-names>Y.</given-names></name> <name><surname>Yin</surname> <given-names>J.</given-names></name> <name><surname>Zheng</surname> <given-names>D.</given-names></name> <name><surname>Hu</surname> <given-names>H.</given-names></name><etal/></person-group> (<year>2023</year>). <article-title>Population genomics and pathotypic evaluation of the bacterial leaf blight pathogen of rice reveals rapid evolutionary dynamics of a plant pathogen.</article-title> <source><italic>Front. Cell Infect. Microbiol.</italic></source> <volume>13</volume>:<fpage>1183416</fpage>. <pub-id pub-id-type="doi">10.3389/fcimb.2023.1183416</pub-id> <pub-id pub-id-type="pmid">37305415</pub-id></citation></ref>
<ref id="B41"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Suzuki</surname> <given-names>S.</given-names></name> <name><surname>Ono</surname> <given-names>N.</given-names></name> <name><surname>Furusawa</surname> <given-names>C.</given-names></name> <name><surname>Ying</surname> <given-names>B. W.</given-names></name> <name><surname>Yomo</surname> <given-names>T.</given-names></name></person-group> (<year>2011</year>). <article-title>Comparison of sequence reads obtained from three next-generation sequencing platforms.</article-title> <source><italic>PLoS One</italic></source> <volume>6</volume>:<fpage>e19534</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0019534</pub-id> <pub-id pub-id-type="pmid">21611185</pub-id></citation></ref>
<ref id="B42"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wilm</surname> <given-names>A.</given-names></name> <name><surname>Aw</surname> <given-names>P. P.</given-names></name> <name><surname>Bertrand</surname> <given-names>D.</given-names></name> <name><surname>Yeo</surname> <given-names>G. H.</given-names></name> <name><surname>Ong</surname> <given-names>S. H.</given-names></name> <name><surname>Wong</surname> <given-names>C. H.</given-names></name><etal/></person-group> (<year>2012</year>). <article-title>LoFreq: A sequence-quality aware, ultra-sensitive variant caller for uncovering cell-population heterogeneity from high-throughput sequencing datasets.</article-title> <source><italic>Nucleic Acids Res.</italic></source> <volume>40</volume> <fpage>11189</fpage>&#x2013;<lpage>11201</lpage>. <pub-id pub-id-type="doi">10.1093/nar/gks918</pub-id> <pub-id pub-id-type="pmid">23066108</pub-id></citation></ref>
<ref id="B43"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zeng</surname> <given-names>M.</given-names></name> <name><surname>Zhou</surname> <given-names>X.</given-names></name> <name><surname>Yang</surname> <given-names>C.</given-names></name> <name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>J.</given-names></name> <name><surname>Xin</surname> <given-names>C.</given-names></name><etal/></person-group> (<year>2023</year>). <article-title>Comparative analysis of the biological characteristics and mechanisms of azole resistance of clinical <italic>Aspergillus fumigatus</italic> strains.</article-title> <source><italic>Front. Microbiol.</italic></source> <volume>14</volume>:<fpage>1253197</fpage>. <pub-id pub-id-type="doi">10.3389/fmicb.2023.1253197</pub-id> <pub-id pub-id-type="pmid">38029222</pub-id></citation></ref>
</ref-list>
</back>
</article>