<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Genet.</journal-id>
<journal-title>Frontiers in Genetics</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Genet.</abbrev-journal-title>
<issn pub-type="epub">1664-8021</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">786825</article-id>
<article-id pub-id-type="doi">10.3389/fgene.2022.786825</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Genetics</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Hybrid Assembly and Annotation of the Genome of the Indian <italic>Punica granatum</italic>, a Superfood</article-title>
<alt-title alt-title-type="left-running-head">Usha et al.</alt-title>
<alt-title alt-title-type="right-running-head">First Draft of Pomegarante Genome <italic>Var Bhagwa</italic>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Usha</surname>
<given-names>Talambedu</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1752062/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Middha</surname>
<given-names>Sushil Kumar</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1348279/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Babu</surname>
<given-names>Dinesh</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/391666/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Goyal</surname>
<given-names>Arvind Kumar</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1517825/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Das</surname>
<given-names>Anupam J.</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Saini</surname>
<given-names>Deepti</given-names>
</name>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1752083/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Sarangi</surname>
<given-names>Aditya</given-names>
</name>
<xref ref-type="aff" rid="aff7">
<sup>7</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Krishnamurthy</surname>
<given-names>Venkatesh</given-names>
</name>
<xref ref-type="aff" rid="aff8">
<sup>8</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1505471/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Prasannakumar</surname>
<given-names>Mothukapalli Krishnareddy</given-names>
</name>
<xref ref-type="aff" rid="aff9">
<sup>9</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/685798/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Saini</surname>
<given-names>Deepak Kumar</given-names>
</name>
<xref ref-type="aff" rid="aff10">
<sup>10</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Sidhalinghamurthy</surname>
<given-names>Kora Rudraiah</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Department of Biochemistry</institution>, <institution>Bangalore University</institution>, <addr-line>Bengaluru</addr-line>, <country>India</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>DBT-BIF Facility</institution>, <institution>Department of Biotechnology</institution>, <institution>Maharani Lakshmi Ammanni College for Women</institution>, <addr-line>Bengaluru</addr-line>, <country>India</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Faculty of Pharmacy and Pharmaceutical Sciences</institution>, <institution>University of Alberta</institution>, <addr-line>Edmonton</addr-line>, <addr-line>AB</addr-line>, <country>Canada</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Centre for Bamboo Studies</institution>, <institution>Department of Biotechnology</institution>, <institution>Bodoland University</institution>, <addr-line>Kokrajhar</addr-line>, <country>India</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Molsys Pvt. Ltd.</institution>, <addr-line>Bengaluru</addr-line>, <country>India</country>
</aff>
<aff id="aff6">
<sup>6</sup>
<institution>Protein Design Private Limited</institution>, <addr-line>Bengaluru</addr-line>, <country>India</country>
</aff>
<aff id="aff7">
<sup>7</sup>
<institution>Basesolve Informatics Private Limited</institution>, <addr-line>Ahmedabad</addr-line>, <country>India</country>
</aff>
<aff id="aff8">
<sup>8</sup>
<institution>Genotypic Technology Pvt Limited</institution>, <addr-line>Bengaluru</addr-line>, <country>India</country>
</aff>
<aff id="aff9">
<sup>9</sup>
<institution>Department of Plant Pathology</institution>, <institution>University of Agricultural Sciences</institution>, <addr-line>Bengaluru</addr-line>, <country>India</country>
</aff>
<aff id="aff10">
<sup>10</sup>
<institution>Department of Molecular Reproduction Development and Genetics</institution>, <institution>Indian Institute of Science</institution>, <addr-line>Bengaluru</addr-line>, <country>India</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/490014/overview">Sunil Kumar Sahu</ext-link>, Beijing Genomics Institute (BGI), China</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1308287/overview">Yu Zhang</ext-link>, Sun Yat-sen University, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1499461/overview">Anna Yssel</ext-link>, University of Cape Town, South Africa</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Sushil Kumar Middha, <email>drsushilmiddha@mlacw.edu.in</email>; Deepak Kumar Saini, <email>deepaksaini@iisc.ac.in</email>
</corresp>
<fn fn-type="other">
<p>This article was submitted to Plant Genomics, a section of the journal Frontiers in Genetics</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>11</day>
<month>05</month>
<year>2022</year>
</pub-date>
<pub-date pub-type="collection">
<year>2022</year>
</pub-date>
<volume>13</volume>
<elocation-id>786825</elocation-id>
<history>
<date date-type="received">
<day>05</day>
<month>10</month>
<year>2021</year>
</date>
<date date-type="accepted">
<day>15</day>
<month>03</month>
<year>2022</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2022 Usha, Middha, Babu, Goyal, Das, Saini, Sarangi, Krishnamurthy, Prasannakumar, Saini and Sidhalinghamurthy.</copyright-statement>
<copyright-year>2022</copyright-year>
<copyright-holder>Usha, Middha, Babu, Goyal, Das, Saini, Sarangi, Krishnamurthy, Prasannakumar, Saini and Sidhalinghamurthy</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>The wonder fruit pomegranate (<italic>Punica granatum</italic>, family Lythraceae) is one of India&#x2019;s economically important fruit crops that can grow in different agro-climatic conditions ranging from tropical to temperate regions. This study reports high-quality <italic>de novo</italic> draft hybrid genome assembly of diploid Punica cultivar &#x201c;Bhagwa&#x201d; and identifies its genomic features. This cultivar is most common among the farmers due to its high sustainability, glossy red color, soft seed, and nutraceutical properties with high market value. The draft genome assembly is about 361.76&#xa0;Mb (N50 &#x3d; 40&#xa0;Mb), &#x223c;9.0&#xa0;Mb more than the genome size estimated by flow cytometry. The genome is 90.9% complete, and only 26.68% of the genome is occupied by transposable elements and has a relative abundance of 369.93 SSRs/Mb of the genome. A total of 30,803 proteins and their putative functions were predicted. Comparative whole-genome analysis revealed <italic>Eucalyptus grandis</italic> as the nearest neighbor. KEGG-KASS annotations indicated an abundance of genes involved in the biosynthesis of flavonoids, phenylpropanoids, and secondary metabolites, which are responsible for various medicinal properties of pomegranate, including anticancer, antihyperglycemic, antioxidant, and anti-inflammatory activities. The genome and gene annotations provide new insights into the pharmacological properties of the secondary metabolites synthesized in pomegranate. They will also serve as a valuable resource in mining biosynthetic pathways for key metabolites, novel genes, and variations associated with disease resistance, which can facilitate the breeding of new varieties with high yield and superior quality.</p>
</abstract>
<kwd-group>
<kwd>Punica granatum (cultivar Bhagwa)</kwd>
<kwd>flavonoids biosynthesis</kwd>
<kwd>phenylpropanoid pathway</kwd>
<kwd>whole genome</kwd>
<kwd>oxford nanopore</kwd>
<kwd>hybrid assembly</kwd>
<kwd>next-generation sequencing</kwd>
</kwd-group>
</article-meta>
</front>
<body>
<sec id="s1">
<title>Introduction</title>
<p>
<italic>Punica granatum</italic> L (family: Lythraceae), alias pomegranate, is one of the ancient and well-known edible fruits and is well known in Ayurveda as a Rasayana (life- and health-span enhancing agent) (<xref ref-type="bibr" rid="B4">Balasubramani et al., 2014</xref>). The genus <italic>Punica</italic> (Angiosperm Phylogeny Group IV classification) contains only two sister species with a classic intercontinental disjunction dispersion, one in Western Asia, Iran (<italic>P. granatum</italic>), and the other in Socotra Island, Yemen (<italic>P. protopunica</italic>) (<xref ref-type="bibr" rid="B23">Kandylis and Kokkinomagoulos, 2020</xref>). <italic>P. granatum</italic> is culturally considered a symbol of fertility, abundance, blessings, immortality, and invincibility because of its pharmaceutical and nutraceutical values (<xref ref-type="bibr" rid="B45">Usha et al., 2020</xref>). The plant is domesticated in Asia, the Middle East, Southern Europe, the United States, and the milder climatic regions of Africa for food, religious, and medicinal uses. Interestingly, every part of the pomegranate, namely, the fruit, rind, flowers, leaves, roots, and wood, has therapeutic and economic values (<xref ref-type="bibr" rid="B31">Medjakovic and Jungbauer, 2013</xref>). The chemical constituents of the fruits vary based on the cultivar, growth climate, maturity time, cultivation method, and storage conditions (<xref ref-type="bibr" rid="B16">Fadavi et al., 2005</xref>). The peel, seed, bark, leaves, seed oil, juice, and heartwood of pomegranate contain several potentially active phytochemicals such as alkaloids, anthocyanins, flavonoids, gallotannins, organic acids, polyphenols, proanthocyanidins, tannins, terpenes, tocopherols, conjugated linolenic acids, triacylglycerols, sterols, steroids, minerals, and complex polysaccharides (<xref ref-type="bibr" rid="B43">Sreekumar et al., 2014</xref>; <xref ref-type="bibr" rid="B45">Usha et al., 2020</xref>). There is ample evidence demonstrating the therapeutic effects of pomegranate and its derived products in arthritis, bacterial infections, diabetes, dental conditions, cardiac disorders, erectile dysfunction, hyperlipidemia, Alzheimer&#x2019;s, infant brain ischemia, obesity (<xref ref-type="bibr" rid="B22">Jurenka, 2008</xref>), cancer (breast, skin, prostate, colon, thyroid, and osteosarcoma) (<xref ref-type="bibr" rid="B45">Usha et al., 2020</xref>), AIDS (<xref ref-type="bibr" rid="B43">Sreekumar et al., 2014</xref>), and inflammation (<xref ref-type="bibr" rid="B46">Vu&#x10d;i&#x107; et al., 2005</xref>). Moreover, 73 clinical trials have been conducted to date, exploring the efficacy and safety of pomegranate mono- and polyherbal medicines in a wide array of ailments<xref ref-type="fn" rid="fn1">
<sup>1</sup>
</xref>.</p>
<p>Pomegranate is thus categorized as a superfood, and there is a soaring demand for its fruit, processed products, and byproducts. In addition to its therapeutic value, the unique morphological characteristics, namely, 1) andromonoecy (<xref ref-type="bibr" rid="B25">Lazare et al., 2020</xref>), 2) each aril being derived from a single ovule accompanied by independent fertilization, and 3) edible juicy, fleshy external seed coat encapsulating the inner fibrous seed coat (<xref ref-type="bibr" rid="B38">Qin et al., 2017</xref>), make <italic>P. granatum</italic> a fascinating fruit to study reproduction, selective adaptation, and evolution in plants.</p>
<p>The pomegranate plant is unique in that it can thrive in various agro-climatic conditions, from tropical to temperate, which are typically deemed unsuitable for cultivating many other economically and medicinally important fruits (<xref ref-type="bibr" rid="B9">Chandra et al., 2010</xref>). Due to its high therapeutic value and global demand, there has been a tremendous increase in export potential in recent years, with notably India being one of the largest exporters of pomegranate to the world. Of the ten available cultivars in India, &#x201c;Bhagwa&#x201d; is the sustained variety that is primarily exported and the most popular among the farmers (<xref ref-type="bibr" rid="B9">Chandra et al., 2010</xref>). Thus, there is a compelling need to use modern molecular genetics methods to obtain insights into this cultivar&#x2019;s genetic and molecular features, aimed at producing high-quality pomegranate fruits with an attractive appearance and a relatively high content of health-promoting ingredients, and disease resistance. Three currently available genome sequences of the Chinese varieties &#x201c;Dabenzi&#x201d; (328.38&#xa0;Mb) (<xref ref-type="bibr" rid="B38">Qin et al., 2017</xref>), &#x201c;Taishanhong&#x201d; (274&#xa0;Mb) (<xref ref-type="bibr" rid="B50">Yuan et al., 2018</xref>), and &#x201c;Tunisia&#x201d; (320.31&#xa0;Mb) (<xref ref-type="bibr" rid="B28">Luo et al., 2020</xref>) and their gene annotations have facilitated advances in basic research, comparative, and evolutionary genomics studies of <italic>P. granatum.</italic> While these resources help dissect the metabolic features, the sequence for the common Indian cultivar must be obtained for trait improvement and to further enhance the production of secondary metabolites in the Indian variety &#x201c;Bhagwa.&#x201d;</p>
<p>The current study presents the first <italic>de novo</italic> draft genome, hybrid assembly of the <italic>P. granatum</italic> of the Indian soft-seeded variety &#x201c;Bhagwa&#x201d; by Illumina and Oxford Nanopore sequencing technologies. We have identified a large set of genes involved in the production of secondary metabolites with medicinal values, such as phenylpropanoids, flavonoids, and tocopherols. The Hidden Gene prediction model unraveled the similarity of the <italic>P. granatum</italic> genome to the <italic>Eucalyptus grandis.</italic> The repeat elements and microsatellites in the assembled genome occupied only 26.68% and 0.06% of the genome. Our new findings of this draft genome will help understand the metabolic traits and improve the quality of the &#x201c;Bhagwa&#x201d; variety and facilitate approaches for increasing the content of secondary metabolites. This first draft whole-genome sequence of an Indian cultivar also presents an essential template for comparative genome analysis for crops from different geographic regions.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>Materials and Methods</title>
<p>The complete workflow adopted in the study is provided in <xref ref-type="fig" rid="F1">Figure 1</xref>.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Overview of the study from the isolation of DNA to assembly and analysis of the genome.</p>
</caption>
<graphic xlink:href="fgene-13-786825-g001.tif"/>
</fig>
<sec id="s2-1">
<title>Estimation of Nuclear DNA Content of the Leaf From <italic>Punica granatum</italic> &#x201c;Bhagwa&#x201d; Variety by Flow Cytometry</title>
<p>Approximately 1&#xa0;cm<sup>2</sup> of young <italic>P. granatum</italic> leaf tissue was taken and nuclear suspensions were made according to <xref ref-type="bibr" rid="B20">Galbraith et al. (1983)</xref>, with slight changes. The leaves of a sample were chopped with a razor blade in a Petri dish with 1&#xa0;ml cold Galbraith nuclear isolation buffer: 45&#xa0;mM magnesium chloride (Merck, Germany), 30&#xa0;mM tri-sodium citrate, 20&#xa0;mM 4-morpholinepropanesulfonic acid, 3-propanesulfonic acid (MOPS) pH 7, 10&#xa0;mM sodium metabisulfite, 1% polyvinylpyrrolidone 10,000, 1% (w/v), and Triton X-100 (Sigma-Aldrich, United States). The nuclear suspension was filtered using a nylon mesh of 50&#xa0;&#xb5;m to eliminate cell fragments and large debris. Nuclei were stained with 50&#xa0;&#x3bc;g&#xa0;ml<sup>&#x2212;1</sup> propidium iodide (PI) (Sigma-Aldrich, United States), a DNA-intercalating fluorochrome, and 50&#xa0;&#x3bc;g&#xa0;ml<sup>&#x2212;1</sup> RNase A (Sigma, United States) was also added. The samples were then incubated in the dark for 15&#xa0;min on ice before flow cytometric evaluation. Chicken erythrocyte nuclei were used as a standard for nuclear content estimation. The stained standard and plant nuclei (5,000&#x2013;10,000 events) were then analyzed on a BD Accuri C6 flow cytometer at 488&#xa0;nm, and the fluorescence signal was collected using a 585/40&#xa0;nm bandpass filter.</p>
</sec>
<sec id="s2-2">
<title>Isolation of Genomic DNA From the Leaves of the <italic>Punica granatum</italic> &#x201c;Bhagwa&#x201d; Variety</title>
<p>The fresh leaves from a sapling of <italic>P. granatum</italic>, collected from Gandhi Krishi Vigyana Kendra (GKVK), Bengaluru, India, were harvested for genome sequencing. &#x201c;Bhagwa&#x201d; was chosen for genome sequencing due to its extensive domestication in India. The fresh leaves were frozen immediately in liquid nitrogen to extract high molecular weight genomic DNA using the cetyltrimethylammonium bromide (CTAB) method. 1&#xa0;gm of <italic>P. granatum</italic> leaves were weighed and frozen in liquid nitrogen and homogenized at 4,000&#xa0;rpm for 10&#xa0;s, using TOMY Micro Smash MS 100 cell disruptor. The homogenization step was repeated twice; the thus obtained lysate was mixed with 1&#xa0;ml CTAB lysis buffer (100&#xa0;mM Tris base, 2% CTAB, 1.5&#xa0;M NaCl, 20&#xa0;mM EDTA, 1% sodium dodecyl sulphate (SDS), and 1% polyethylene glycol (Sigma, United States) and incubated for 30&#xa0;min at 65&#xb0;C. 20&#xa0;&#xb5;l of proteinase K solution was added to the vial and incubated for 1&#xa0;h at 56&#xb0;C.</p>
<p>Further, proteinase K (Sigma, United States) activity was arrested by incubating at 65&#xb0;C for 10&#xa0;min. The vials were centrifuged at 8,000&#xa0;rpm for 10&#xa0;min at RT to remove the debris. 20&#xa0;&#xb5;l of RNAase (20&#xa0;mg/ml; Sigma, United States) was added to the collected supernatant and incubated at 65&#xb0;C for 10&#xa0;min. To the supernatant, an equal volume of phenol: chloroform: isoamyl alcohol (25:24:1) (Merck, Germany) mixture was added and mixed gently. The mixture was centrifuged at 13,200&#xa0;rpm for 10&#xa0;min at 4&#xb0;C to separate the phases. The aqueous phase from the centrifuged vials was transferred to a fresh vial, and an equal volume of 100% isoamyl-alcohol was added and incubated at &#x2212;80&#xb0;C for 1&#xa0;h. The vials were left at room temperature for 5&#xa0;min and centrifuged at 13,200&#xa0;rpm for 20&#xa0;min at 4&#xb0;C. The pellet was washed twice with 70% ethanol. The extracted DNA was run on 1% agarose gel to assess the quality. DNA was also analyzed for its concentration and purity using NanoDrop&#x2122; Spectrophotometer and Qubit 4 Fluorometer (Thermo Fisher Scientific, Massachusetts, United States). 1x TE buffer was added to the pellets of the two vials for Oxford Nanopore library preparation. The pellets in the other two vials were resuspended in &#xd7;10 Tris buffer (pH: 8) and were further used for Illumina library preparation.</p>
</sec>
<sec id="s2-3">
<title>Library Preparation Methods and Whole-Genome Sequencing</title>
<p>The extracted DNA was purified using Qiagen DNeasy Blood &#x26; Tissue kit column (Qiagen, Germany). Paired-end (PE) and mate-pair (MP) sequencing libraries with insert sizes of 400&#x2013;550&#x2009;bp and 300&#xa0;bp to 1,000&#xa0;bp, respectively, were constructed and sequenced on the Illumina HiSeq 2,000 platform to obtain low error short reads. For nanopore sequencing, 2&#xa0;&#xb5;g of genomic DNA was end-repaired (NEBNext Ultra II End Repair Kit, New England Biolabs, MA, United States) and cleaned with &#xd7;1 AMPure beads (Beckman Coulter, United States). NEB blunt/TA ligase (New England Biolabs, MA, United States) was used to perform adapter ligations (AMX) for 30&#xa0;min. Library mix was cleaned up using 0.4&#xd7; AMPure beads (Beckman Coulter, United States) and finally eluted in 16&#xa0;&#xb5;l of elution buffer. A total of 480&#xa0;ng of sequencing library was obtained and used for sequencing. Long reads were obtained by sequencing on MinION MklB (Oxford Nanopore Technologies, Oxford, United Kingdom) using spot on flow cell (R9.4), and base calling was performed using Metrichor Nanopore. A total of 46.2&#xa0;Gb of raw data was generated on the Oxford Nanopore and Illumina platforms.</p>
</sec>
<sec id="s2-4">
<title>Raw Data Processing and <italic>In Silico</italic> Genome Size Estimation</title>
<p>The Illumina paired-end and mate-pair raw reads were checked for quality using FastQC (<xref ref-type="bibr" rid="B3">Andrews, 2010</xref>). FASTP (<xref ref-type="bibr" rid="B10">Chen et al., 2018</xref>) was used to process Illumina raw reads for adapters and trim low-quality bases. The raw nanopore reads were processed, and the quality was checked by LongQC (<xref ref-type="bibr" rid="B19">Fukasawa et al., 2020</xref>), followed by quality trimming using Filtlong<xref ref-type="fn" rid="fn8">
<sup>2</sup>
</xref>. The processed nanopore reads were corrected using NECAT (<xref ref-type="bibr" rid="B11">Chen et al., 2021</xref>). The genome size estimation of <italic>P. granatum</italic> was carried out using GenomeScope V2.0. (<xref ref-type="bibr" rid="B47">Vurture et al., 2017</xref>).</p>
</sec>
<sec id="s2-5">
<title>
<italic>De Novo</italic> Hybrid Assembly of the Nuclear Genome of <italic>P. granatum</italic>
</title>
<p>The <italic>P. granatum</italic> draft genome was assembled using MaSuRCA V.3.4.2<xref ref-type="fn" rid="fn3">
<sup>3</sup>
</xref> (<xref ref-type="bibr" rid="B52">Zimin et al., 2013</xref>) and WENGAN (<xref ref-type="bibr" rid="B15">Di Genova et al., 2021</xref>) hybrid assemblers individually based on the Illumina paired-end reads and corrected nanopore reads. The draft assemblies generated by MaSuRCA and WENGAN were merged. Prokaryotic contamination was removed from the merged assembly using EukRep (<xref ref-type="bibr" rid="B48">West et al., 2018</xref>), followed by developing one non-redundant set of contigs by assembly reduction using Redundans (<xref ref-type="bibr" rid="B37">Pryszcz and Gabald&#xf3;n, 2016</xref>). The reduced contigs were further scaffolded based on <italic>P. granatum</italic> reference (GCF_007655135.1) using RagTag (<xref ref-type="bibr" rid="B2">Alonge et al., 2019</xref>).</p>
</sec>
<sec id="s2-6">
<title>Qualitative Analysis of <italic>De Novo P. granatum</italic> Draft Genome Assembly</title>
<p>Benchmarking set of Universal Single-Copy Orthologues (BUSCO-version 5) (<xref ref-type="bibr" rid="B41">Sim&#xe3;o et al., 2015</xref>) was used for the identification of single-copy orthologs against eudicot odb10 lineage in the <italic>P. granatum</italic> draft assembly and compared with <italic>P. granatum</italic> cultivarDabenzi [GCA_002201585.1], Taishanhong [GCA_002864125.1], isolate Tunisia 2019 [GCF_007655135.1], and strain AG 2017 [GCA_002837095.1]. BUSCO robustly estimates the completeness of the genome, and BUSCOs are conserved single-copy orthologs that are predicted to be present in the complete genome. Therefore, the number of fragmented, duplicated, present, and missing BUSCOs can be used in the quality control of the assembled genome.</p>
</sec>
<sec id="s2-7">
<title>Nuclear Genome Annotations</title>
<sec id="s2-7-1">
<title>Identification of Transposable Elements</title>
<p>Transposable elements (TEs) are major players in the structure and evolution of plant genomes. Their ability to locomote around and replicate within genomes are probably the most essential contributors to genome size and plasticity (<xref ref-type="bibr" rid="B39">Sahebi et al., 2018</xref>). Thorough annotation of TEs is ideal for dealing with the deluge of genome data. Therefore, <italic>de novo</italic> repeat identification was performed by Repeat Modeller2 (<xref ref-type="bibr" rid="B18">Flynn et al., 2020</xref>), LongRepMarker (<xref ref-type="bibr" rid="B27">Liao et al., 2021</xref>), and Extensive De Novo TE Annotator (EDTA) (<xref ref-type="bibr" rid="B35">Ou et al., 2019</xref>). The unclassified repeats were further classified using DeepTE (<xref ref-type="bibr" rid="B49">Yan et al., 2020</xref>). The <italic>de novo</italic> repeats identified by all the above three methods were merged and clustered using cd-hit-est with a 99% threshold to generate one non-redundant repeat library. The draft genome assembly was used to unravel the transpositional landscape of <italic>P. granatum</italic> using RepeatMasker 4.0.7<xref ref-type="fn" rid="fn4">
<sup>4</sup>
</xref> against the <italic>de novo</italic> constructed repeat library.</p>
</sec>
<sec id="s2-8">
<title>Prediction of Microsatellites</title>
<p>The plant genomes are filled with low complexity sequences such as simple sequence repeats (SSRs). The MISA tool<xref ref-type="fn" rid="fn2">
<sup>5</sup>
</xref> was used to screen for the presence of SSRs in scaffolds. The sequences were annotated and screened for the most frequent type of SSR motif families and mono-repeats recurring a minimum of 10 times, di-repeats recurring a minimum of 5 times, and tri/tetra/penta/hexa-repeats recurring a minimum of 5 times. SSR statistics were generated by PySSRstat<xref ref-type="fn" rid="fn7">
<sup>6</sup>
</xref>.</p>
</sec>
<sec id="s2-8-1">
<title>Gene Annotations</title>
<p>Genome annotation was done using Modular Open Source Genome Annotator MOSGA (<xref ref-type="bibr" rid="B30">Martin et al., 2021</xref>). BRAKER2 (<xref ref-type="bibr" rid="B6">Br&#x16f;na et al., 2021</xref>) was used for gene prediction, with orthology-based-evidence mode using OrthoDB as a data source that relies on GeneMark-EP spaln and DIAMOND (<xref ref-type="bibr" rid="B7">Buchfink et al., 2015</xref>). Prediction of tRNA sequences was performed using tRNAscan-SE2.0 (<xref ref-type="bibr" rid="B8">Chan et al., 2021</xref>). Prediction of rRNAs was done using Barrnap<xref ref-type="fn" rid="fn9">
<sup>7</sup>
</xref>. Functional gene prediction was made by comparison against Swiss-Prot and EggNog 5 protein databases. The gene model was obtained in GFF format. The transcript sequences were extracted from the GFF file using the gffread utility<xref ref-type="fn" rid="fn1">
<sup>8</sup>
</xref>. The transcript FASTA sequences were subjected to blastx against the Eudicotyledons (taxonomy id: 22663) of NCBI non-redundant database using DIAMOND [parameters: -max-target-seqs 20 --outfmt 5 --sensitive -e 1e-5 -b12 -c1 --taxonlist 22,663]. The BLASTX outputs in XML format were further annotated using BLAST2GO (<xref ref-type="bibr" rid="B13">Conesa et al., 2005</xref>).</p>
</sec>
<sec id="s2-8-2">
<title>Comparative Genomic Analysis Between Pomegranate and Other Plant Species</title>
<p>A gene family cluster analysis of the complete gene sets of pomegranate (<italic>P. granatum</italic>), <italic>E. grandis</italic>, apple (<italic>Malus domestica</italic>), arabidopsis (<italic>Arabidopsis thaliana</italic>), and grape (<italic>Vitis vinifera</italic>) was performed using OrthoVenn2 (<ext-link ext-link-type="uri" xlink:href="https://orthovenn2.bioinfotoolkits.net/">https://orthovenn2.bioinfotoolkits.net/</ext-link>).</p>
</sec>
<sec id="s2-8-3">
<title>Phylogenetic Tree Based on Average Nucleotide Identity Analysis</title>
<p>We calculated the average nucleotide identity between the <italic>P. granatum</italic> draft genome with <italic>Prunus</italic> to check the genetic relatednes<italic>s. persica</italic> (cultivar_Lovell), <italic>P. granatum</italic> (cultivar Bhagwa, Dabenzi, Taishanhong, Tunisia 2019, and strain AG 2017), <italic>Sonneratia. alba</italic>, <italic>Vitis. vinifera</italic> (cultivar PN40024), <italic>Sonneratia. caseolaris</italic>, <italic>Eucalyptus. grandis</italic> (isolate ANBG69807), <italic>Corymbia. maculata</italic> (isolate sf003), <italic>Angophora. floribunda</italic> (isolate sf002), and <italic>Corymbia citriodora</italic> (subsp. variegata) on the &#x201c;pyANI&#x201d; software.</p>
</sec>
<sec id="s2-8-4">
<title>Identification and Analysis of the Variants in <italic>P. granatum</italic> Genome</title>
<p>The trimmed Illumina paired-end reads were aligned using bowtie2<xref ref-type="fn" rid="fn10">
<sup>9</sup>
</xref> against the representative genome reference of <italic>P. granatum</italic> obtained from NCBI (GCF_007655135.1_ASM765513v2). The aligned reads were sorted and indexed using Samtools (<xref ref-type="bibr" rid="B26">Li et al., 2009</xref>). The aligned reads were utilized in variant calling using Genome Analysis Toolkit (GATK) (<xref ref-type="bibr" rid="B14">DePristo et al., 2011</xref>) with optimal practices for germline SNPs and indel calling<xref ref-type="fn" rid="fn6">
<sup>10</sup>
</xref> (<xref ref-type="bibr" rid="B36">Pramesh et al., 2020</xref>). At first, the HaplotypeCaller was called without Base Quality Score Recalibration (BQSR), the thus obtained variants were fed to BQSR as know variants, and a recalibration was achieved, while no other modifications were made in the workflow. A set of hard filters such as &#x201c;QD &#x3c; 2.0 &#x7c;&#x7c; FS &#x3e; 200.0 &#x7c;&#x7c; ReadPosRankSum &#x3c; &#x2212;20.0 &#x7c;&#x7c; InbreedingCoeff &#x3c; &#x2212;0.8&#x201d;, &#x201c;SB &#x3e;&#x3d; 0.10 &#x7c;&#x7c; QD &#x3c; 5.0 &#x7c;&#x7c; HRun &#x3e;&#x3d; 4&#x201d; and other parameters were applied, such as cluster size 3, mask extension 5, and cluster window size 10. Thus, the filtered variants were predicted using snpEFF (<xref ref-type="bibr" rid="B12">Cingolani et al., 2012</xref>) with a custom-built effects database (<ext-link ext-link-type="uri" xlink:href="http://snpeff.sourceforge.net/SnpEff_manual.html">http://snpeff.sourceforge.net/SnpEff_manual.html&#x23;databases</ext-link>) for <italic>P. granatum.</italic>
</p>
</sec>
</sec>
</sec>
<sec sec-type="results" id="s3">
<title>Results</title>
<sec id="s3-1">
<title>Sequence Assembly</title>
<p>Flow cytometric analysis projected the <italic>P. granatum</italic> genome to be around 352.38&#xa0;Mb (<xref ref-type="sec" rid="s11">Supplementary Figure S1</xref>). The genome was sequenced using two sequencing platforms: Illumina (short-read technology) and Oxford Nanopore (long-read technology). We generated 52&#xa0;Gb paired reads (2&#x2a;150&#xa0;bp), 11&#xa0;GB mate-pair reads, and 2&#xa0;Gb long reads of 2.31956&#xa0;Kb average length. A total of 65&#xa0;Gb of data representing &#xd7;155 fold coverage was generated (<xref ref-type="sec" rid="s11">Supplementary Table S1</xref>). Data obtained through Illumina sequencing technology was also used to estimate the genome size and was found to be 276.60&#xa0;Mb. (<xref ref-type="sec" rid="s11">Supplementary Figure S2</xref>).</p>
<p>The raw reads (Bioproject: PRJNA407279) were filtered based on the quality score, trimmed (<xref ref-type="sec" rid="s11">Supplementary Table S1</xref>), and then assembled into contigs using assemblers MaSuRCA V3.4.2 and WENGAN. The aforementioned assembler statistics resulted in 122 contigs of maximum length up to 88134655 bp with an N50 of 40&#xa0;Mb, L50 of 3, and a hybrid assembly of 361.76&#xa0;Mb of <italic>P. granatum</italic> (<xref ref-type="table" rid="T1">Table 1</xref>). The assembled genome was made up of 38.86% of GC.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Statistics of draft genome assembly of <italic>P. granatum</italic>.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Statistics without reference</th>
<th align="center">
<italic>P. granatum</italic> references (GCF_007655135)</th>
<th align="center">
<italic>P. granatum</italic> draft assembly</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Estimated genome size (Mb)</td>
<td align="center">313.18&#xa0;Mb</td>
<td align="center">352.38&#xa0;Mb</td>
</tr>
<tr>
<td align="left">Assembled genome size (Mb)</td>
<td align="center">320.31&#xa0;Mb</td>
<td align="center">361.76&#xa0;Mb</td>
</tr>
<tr>
<td align="left">Number of scaffolds (&#x2265;1&#xa0;kb)</td>
<td align="center">473</td>
<td align="center">122</td>
</tr>
<tr>
<td align="left">N50 scaffold length (Mb)</td>
<td align="center">39.02</td>
<td align="center">39.79</td>
</tr>
<tr>
<td align="left">Longest scaffold (Mb)</td>
<td align="center">54.25</td>
<td align="center">86.06</td>
</tr>
<tr>
<td align="left">Total size of assembled contigs (Mb)</td>
<td align="center">320.33&#xa0;Mb</td>
<td align="center">361.76&#xa0;Mb</td>
</tr>
<tr>
<td align="left">Number of contigs (&#x2265;1&#xa0;kb)</td>
<td align="center">661</td>
<td align="center">7,640</td>
</tr>
<tr>
<td align="left">Number of contigs (&#x2265; 50,000 bp)</td>
<td align="center">399</td>
<td align="center">1,555</td>
</tr>
<tr>
<td align="left">N50 contig length (kb)</td>
<td align="center">4,489.929</td>
<td align="center">61.031</td>
</tr>
<tr>
<td align="left">Largest contig (kb)</td>
<td align="center">14,772.832&#xa0;kb</td>
<td align="center">487.599&#xa0;kb</td>
</tr>
<tr>
<td align="left">Total length</td>
<td align="center">320494280</td>
<td align="center">361760465</td>
</tr>
<tr>
<td align="left">N50 length (Mb)</td>
<td align="center">39.95&#xa0;Mb</td>
<td align="center">40.75&#xa0;Mb</td>
</tr>
<tr>
<td align="left">GC (%)</td>
<td align="center">40.38</td>
<td align="center">38.86</td>
</tr>
<tr>
<td align="left">Number of genes</td>
<td align="center">33,594</td>
<td align="center">30,803</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-2">
<title>BUSCO Analysis</title>
<p>The quantitative assessment of <italic>P. granatum</italic> draft <italic>de novo</italic> hybrid genome assembly and annotation completeness carried out using BUSCO is represented in <xref ref-type="table" rid="T2">Table 2</xref>. The draft assembly consists of 90.6% of complete BUSCOs and a very minute percentage of missing, fragmented, and duplicated BUSCOs, indicating the completeness of assembled genome. Further, the draft genome completeness is comparable to that of all other assemblies of <italic>P. granatum</italic> available at NCBI.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Qualitative analysis of draft genome assembly.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Measures</th>
<th align="center">
<italic>P. granatum</italic> draft assembly</th>
<th align="center">
<italic>P. granatum</italic> cultivar Dabenzi</th>
<th align="center">
<italic>P. granatum</italic> cultivar Taishanhong</th>
<th align="center">
<italic>P. granatum</italic> isolate Tunisia 2019</th>
<th align="center">
<italic>P. granatum</italic> strain AG2017</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">No. (percentage) of complete BUSCOs (C)</td>
<td align="center">2,114 (90.9%)</td>
<td align="center">2,156 (92.7%)</td>
<td align="center">2,159 (92.8%)</td>
<td align="center">2,150 (92.5%)</td>
<td align="center">2,114 (90.9%)</td>
</tr>
<tr>
<td align="left">No. (percentage) of complete and single-copy BUSCOs (S)</td>
<td align="center">2,012 (86.5%)</td>
<td align="center">2,096 (90.1%)</td>
<td align="center">2,099 (90.2%)</td>
<td align="center">2,069 (89.0%)</td>
<td align="center">2,062 (88.7%)</td>
</tr>
<tr>
<td align="left">No. (percentage) of complete and duplicated BUSCOs (D)</td>
<td align="center">102 (4.4%)</td>
<td align="center">60 (2.6%)</td>
<td align="center">60 (2.6%)</td>
<td align="center">81 (3.5%)</td>
<td align="center">52 (2.2%)</td>
</tr>
<tr>
<td align="left">No. (percentage) of fragmented BUSCOs (F)</td>
<td align="center">91 (3.9%)</td>
<td align="center">79 (3.4%)</td>
<td align="center">76 (3.3%)</td>
<td align="center">84 (3.6%)</td>
<td align="center">93 (4.0%)</td>
</tr>
<tr>
<td align="left">No. (percentage) of missing BUSCOs (M)</td>
<td align="center">121 (5.2%)</td>
<td align="center">91 (3.9%)</td>
<td align="center">91 (3.9%)</td>
<td align="center">92 (3.9%)</td>
<td align="center">119 (5.1%)</td>
</tr>
<tr>
<td align="left">Total BUSCO groups searched</td>
<td align="center">2,326</td>
<td align="center">2,326</td>
<td align="center">2,326</td>
<td align="center">2,326</td>
<td align="center">2,326</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>BUSCO (Benchmarking set of universal single-copy orthologues) result for draft assembly, Dabenzi, Taishanhong, isolate Tunisia 2019, and strain AG2017 of P. granatum.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<sec id="s3-2-1">
<title>Transposable Elements in the Genome</title>
<p>The whole-genome data were also used to identify the transpositional landscape of pomegranate, represented in <xref ref-type="table" rid="T3">Table 3</xref>. <italic>P. granatum</italic> genome harbors both Type I and Type II transposable element (TE) families. The repeat elements occupy only 26.68% of the genome. Long terminal repeat (LTR) retroelements are the most abundant TEs, representing 11.43% of the assembly. Among the LTR elements, Copia/TY1 are more copious than Gypsy/DIRS1.</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Types of transposable elements identified in <italic>P. granatum</italic> genome. The total interspresed repeats mentioned at the bottom of the table 26.68 is the total of retroelements (13.69), DNA transposons (11.23) and unclassified (1.76) Hence these values are highlighted.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Types of Transposable element</th>
<th align="center">Number of elements</th>
<th align="center">Length occupied in bp</th>
<th align="center">Percentage of sequences</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Retroelements</td>
<td align="center">112,404</td>
<td align="center">49,529,468</td>
<td align="char" char=".">
<bold>13.69</bold>
</td>
</tr>
<tr>
<td align="left">SINEs</td>
<td align="center">13,048</td>
<td align="center">1,498,787</td>
<td align="char" char=".">0.41</td>
</tr>
<tr>
<td align="left">LINEs</td>
<td align="center">25,893</td>
<td align="center">6,693,226</td>
<td align="char" char=".">1.85</td>
</tr>
<tr>
<td align="left">(i) CRE/SLACS</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(ii) L2/CR1/Rex</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(iii)R1/LOA/Jockey</td>
<td align="center">21</td>
<td align="center">5,707</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(iv) R2/R4/NeSL</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(v)RTE/Bov-B</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(vi)L1/CIN4</td>
<td align="center">3,843</td>
<td align="center">1,720,126</td>
<td align="char" char=".">0.48</td>
</tr>
<tr>
<td align="left">LTR elements</td>
<td align="center">73,463</td>
<td align="center">41,337,455</td>
<td align="char" char=".">11.43</td>
</tr>
<tr>
<td align="left">(i) BEL/Pao</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(ii) Ty1/Copia</td>
<td align="center">8,660</td>
<td align="center">5,573,403</td>
<td align="char" char=".">1.54</td>
</tr>
<tr>
<td align="left">(iii) Gypsy/DIRS1</td>
<td align="center">3,262</td>
<td align="center">2,504,146</td>
<td align="char" char=".">0.69</td>
</tr>
<tr>
<td align="left">(iv) Retroviral</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">DNA transposons</td>
<td align="center">182,314</td>
<td align="center">40,633,158</td>
<td align="char" char=".">
<bold>11.23</bold>
</td>
</tr>
<tr>
<td align="left">(i) hobo-Activator</td>
<td align="center">683</td>
<td align="center">418,562</td>
<td align="char" char=".">0.12</td>
</tr>
<tr>
<td align="left">(ii) Tc1-IS630-Pogo</td>
<td align="center">24</td>
<td align="center">11,586</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(iii) En-Spm</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(iv) MuDR-IS905</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(v) PiggyBac</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="char" char=".">0</td>
</tr>
<tr>
<td align="left">(vi) Tourist/Harbinger</td>
<td align="center">954</td>
<td align="center">607,183</td>
<td align="char" char=".">0.17</td>
</tr>
<tr>
<td align="left">(vii) Other (Mirage,P-element, Transib)</td>
<td align="center">0</td>
<td align="center">0</td>
<td align="left"/>
</tr>
<tr>
<td align="left">Rolling-circles</td>
<td align="center">835</td>
<td align="center">750,813</td>
<td align="char" char=".">0.21</td>
</tr>
<tr>
<td align="left">(i)Unclassified</td>
<td align="center">30,347</td>
<td align="center">6,356,886</td>
<td align="char" char=".">
<bold>1.76</bold>
</td>
</tr>
<tr>
<td align="left">Total interspersed repeats</td>
<td align="center"/>
<td align="center">96,519,512</td>
<td align="char" char=".">
<bold>26.68</bold>
</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-3">
<title>Microsatellites Predicted in the Genome</title>
<p>A total of 133827 SSRs were revealed in the <italic>P. granatum</italic> genome with a relative abundance of 369.93 SSRs/Mb. The total length of the complete set of microsatellites was 0.06% of the assembled genome. The most frequent motifs were mono- (52.4%) followed by di- (36.2%), tri- (8.8%), tetra- (1.8%), penta- (0.5%), and hexa-nucleotides (0.3%). The genome sequence also yielded both complex and compound SSRs with an overall density of 0.12%. It was also observed that SSRs composed of smaller repeats accounted for a larger percentage, while those with large repeats represented a smaller percentage of SSRs (<xref ref-type="fig" rid="F2">Figure 2A</xref>). The P1, P2, P3, P4, P5, and P6 were divided into two classes based on the SSR length. A total of 25.6% and 75.3% were classified into long and hypervariable class I type SSRs (&#x2265;20 bp) and class II type (10&#x2013;19 bp), respectively (<xref ref-type="fig" rid="F2">Figure 2B</xref>). The distribution of the different SSR types is heterogeneous, particularly in mono- and di-nucleotides. Among P1, T/A (97.6%) were most frequently occurring and were also the most frequent motif in the entire genome, accounting for 51.0%, and G/C (2.4%) were in almost negligible amounts. Among P2, the highly distributed motifs were AT/AT (61.9%) and AG/CT (31.0%). AAT/ATT (40.1%) and AAG/CTT (28.0%) motifs were the most abundant. Among P4, AAAT/ATTT (37.1%), P5 AAAAG/CTTTT (14.2%), AAAAAT/ATTTTT (12.1%), and P6 AAAATC/ATTTTG (34.5%) and AAAAAG/CTTTTT (10.2%) were the most abundant repeats in P3, P4, P5, and P6 classes (<xref ref-type="fig" rid="F2">Figure 2C</xref>).</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Classification and distribution of microsatellites alias SSRs identified in the <italic>P. granatum</italic> genome. <bold>(A)</bold> Proportions of microsatellites with different motif types. P1: mono-nucleotide repeats; P2: di-nucleotide repeats; P3: tri-nucleotide repeats, P4: tetra-nucleotide repeats; P5: penta-nucleotide repeats; p6: hexa-nucleotide repeats; C: complex: no. of SSRs involved in compound formation. <bold>(B)</bold> Percentage of hypervariable class I and variable class II microsatellites in the <italic>P. granatum</italic> genome. <bold>(C)</bold> Frequency of distribution of the most frequently occurring SSR motif families.</p>
</caption>
<graphic xlink:href="fgene-13-786825-g002.tif"/>
</fig>
</sec>
<sec id="s3-4">
<title>Gene Predictions and Annotations</title>
<p>Comprehensive gene predictions and putative functions assignment of <italic>P. granatum</italic> were made by comparing against the eudicotyledons of NCBI NR (non-redundant) database. The obtained hits were further annotated using BLAST2GO, and 30,803 genes were identified. These 30,803 protein sequences were classified into 30 functional classes under three core categories: biological processes (BP), cellular components (CC), and molecular functions (MF), based on homology (<xref ref-type="fig" rid="F3">Figure 3</xref>) (<xref ref-type="sec" rid="s11">Supplementary Table S2</xref>). On closer examination of the species distribution of hits in the UniportKB database, maximum hits were from the <italic>Eucalyptus grandis</italic> plant<italic>.</italic>
</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Bar chart exhibits gene annotations of the functional classes in each of the three major categories, biological process (BP), cellular component (CC), and molecular function (MF), of gene ontology classification.</p>
</caption>
<graphic xlink:href="fgene-13-786825-g003.tif"/>
</fig>
</sec>
</sec>
<sec id="s3-5">
<title>Elucidation of Metabolic Pathways</title>
<p>KEGG Orthology (KO) links directly to known pathways and KO annotations facilitate concurrent pathway identification. Therefore, 30,803 annotated proteins were mapped to 128 reference pathways. All proteins were classified first at the top three levels: cellular component (18,152), molecular function (23,522), and biological process (19,029). The three functional categories are divided into 46 sub-categories corresponding to the KEGG pathways. 274 and 171 KEGG annotated proteins mapped under the category of biosynthesis of secondary metabolites and terpenoids, respectively (<xref ref-type="sec" rid="s11">Supplementary Table S3</xref>). Because various pharmacological properties of pomegranate are attributed to the presence of secondary metabolites produced in different parts of the plant, we sorted 4,446 proteins coding genes for phenylpropanoid and flavonoid biosynthesis pathways.</p>
<sec id="s3-5-1">
<title>Phenylpropanoid Biosynthesis Pathway</title>
<p>The phenylpropanoid biosynthesis pathway (PP) is essential for a plant&#x2019;s growth, development, and defense. It saddles between the primary and the secondary metabolism. The PP starts with phenylalanine, the end product of the shikimate pathway. Phenylalanine ammonia-lyase (PAL) transforms phenylalanine into cinnamic acid, which leads to the formation of p-coumaric acid by the enzymatic action of trans-cinnamate 4-monooxygenase, which then transforms into p-coumaroyl-CoA by p-coumaroyl: CoA ligase (4CL). Both p-coumaric acid and p-coumaroyl-CoA can act as precursors of the lignin monomer pathway. In contrast, p-coumaroyl-CoA is the precursor of the flavonoid, stilbenoid, diarylheptanoid, and gingerol biosynthetic pathways. The current study identified genes that direct the biosynthesis of monolignols and hydroxycinnamic acids, such as ferulic and sinapic acids, and their corresponding esters. Lignin confers pathogen resistance, vascular integrity, and structural support. Genes coding for enzymes involved in the synthesis of all the three monolignols (p-hydroxycinnamyl alcohols): p-coumaryl, coniferyl, and sinapyl alcohols that result in the lignin polymer were identified in abundance. The 17 enzymes identified are PAL; phenylalanine ammonia-lyase [EC:4.3.1.24], 4CL; 4-coumarate--CoA ligase [EC:6.2.1.12], CCR; cinnamoyl-CoA reductase [EC:1.2.1.44], CYP73A; trans-cinnamate 4-monooxygenase [EC:1.14.13.11], caffeic acid 3-O-methyltransferase [EC:2.1.1.68], CYP84A; ferulate-5-hydroxylase [EC:1.14.-.-], caffeoyl-CoA O-methyltransferase [EC:2.1.1.104], shikimate O-hydroxycinnamoyltransferase [EC:2.3.1.133], coumaroylquinate (coumaroylshikimate) 3&#x2032;-monooxygenase [EC:1.14.13.36], cinnamyl-alcohol dehydrogenase [EC:1.1.1.195], peroxidase [EC:1.11.1.7], coniferyl-alcohol glucosyltransferase [EC:2.4.1.111], beta-glucosidase [EC:3.2.1.21], caffeoyl shikimate esterase [EC:3.1.1.-], and coniferyl-aldehyde dehydrogenase [EC:1.2.1.68], serine carboxypeptidase-like 19 [EC:3.4.16.- 2.3.1.91], and eugenol synthase [EC:1.1.1.318] (<xref ref-type="fig" rid="F4">Figure 4</xref>).</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Phenylpropanoid biosynthesis pathway in <italic>P. granatum</italic>. Numbers 1 to 17 represent the enzymes that catalyze the respective reactions. 1) PAL; phenylalanine ammonia-lyase [EC:4.3.1.24]; 2) 4CL; 4-coumarate--CoA ligase [EC:6.2.1.12]; 3) CCR; cinnamoyl-CoA reductase [EC:1.2.1.44]; 4) CYP73A; trans-cinnamate 4-monooxygenase [EC:1.14.13.11]; 5) E2.1.1.68; caffeic acid 3-O-methyltransferase [EC:2.1.1.68]; 6) CYP84A; ferulate-5-hydroxylase [EC:1.14.-.-]; 7) E2.1.1.104; caffeoyl-CoA O-methyltransferase [EC:2.1.1.104]; 8) E2.3.1.133; shikimate O-hydroxycinnamoyltransferase [EC:2.3.1.133]; 9) CYP98A; coumaroylquinate (coumaroylshikimate) 3&#x2032;-monooxygenase [EC:1.14.13.36]; 10) E1.1.1.195; cinnamyl-alcohol dehydrogenase [EC:1.1.1.195]; 11) E1.11.1.7; peroxidase [EC:1.11.1.7]; 12) UGT72E; coniferyl-alcohol glucosyltransferase [EC:2.4.1.111]; 13) bglB; beta-glucosidase [EC:3.2.1.21]; 14) CSE; caffeoylshikimate esterase [EC:3.1.1.-]; 15) REF1; coniferyl-aldehyde dehydrogenase [EC:1.2.1.68]; 16) serine carboxypeptidase-like 19 [EC:3.4.16.- 2.3.1.91]; 17) eugenol synthase [EC:1.1.1.318].</p>
</caption>
<graphic xlink:href="fgene-13-786825-g004.tif"/>
</fig>
</sec>
<sec id="s3-6">
<title>Flavonoid Biosynthesis Pathway</title>
<p>The flavonoids can be classified into six major groups. These compounds impart protection against exposure to ultraviolet (UV) radiation and phytopathogens, help in signaling during nodulation, male fertility, auxin transport, and coloration of flowers, a visual indicator for pollinators. The draft genome of <italic>P. granatum</italic> revealed the presence of 14 genes capable of encoding enzymes for the flavonoid pathway. Like many other plant species, flavonoids are synthesized through the phenylpropanoid pathway in <italic>P. granatum</italic>. Cinnamoyl CoA or p-coumaroyl-CoA serves as a precursor of the flavonoid biosynthesis pathway. The first enzyme specific for the flavonoid pathway, chalcone synthase, produces naringenin chalcone and the formation of stereospecific cyclic naringenin catalyzed by chalcone isomerase. Naringenin is converted to dihydrokaempferol by naringenin 3-dioxygenase. The bifunctional dihydroflavonol-4-reductase catalyzes the reduction of dihydrokaempferol to leucopelargonidin, which is further converted to pelargonidin, an anthocyanidin, by the action of anthocyanidin synthase. The anthocyanidins are reduced to flavan 3-ols (e.g., catechin and epicatechin) by leucoanthocyanidin reductase (LAR) and anthocyanidin reductase (ANR), respectively. The dihydroflavonol (i.e., the dihydrokaempferol) is converted into kaempferol by flavonol synthase. Flavonoid 3&#x2032;,5&#x2032;-hydroxylase catalyzes the formation of quercetin from kaempferol. Flavonoid 3&#x2032;,5&#x2032;-hydroxylase enzyme acts on naringenin, eriodictyol, dihydroquercetin, and dihydrokaempferol. Flavone synthase I (FNS I) or flavone synthase II (FNS II) synthesizes luteolin from naringenin. P-coumaroyl-CoA is converted to feruloyl coA with the help of coumaroylquinate (coumaroylshikimate) 3&#x2032;-monooxygenase and caffeoyl-CoA O-methyltransferase (<xref ref-type="fig" rid="F5">Figure 5</xref>). The presence of genes coding for these proteins in the genome sequence was also confirmed in this pathway analysis.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Flavonoid biosynthetic pathway found in <italic>P. granatum.</italic> Numbers 1 to 14 represent the enzymes that catalyze the respective reactions. 1) CHS; chalcone synthase [EC:2.3.1.74]; 2) E5.5.1.6; chalcone isomerase [EC:5.5.1.6]; 3) E1.14.11.9; naringenin 3-dioxygenase [EC:1.14.11.9]; 4) FLS; flavonol synthase [EC:1.14.11.23]; 5) DFR; bifunctional dihydroflavonol 4-reductase/flavanone 4-reductase [EC:1.1.1.219 1.1.1.234]; 6) E1.14.13.21; flavonoid 3&#x2032;-monooxygenase [EC:1.14.13.21]; 7) ANR; anthocyanidin reductase [EC:1.3.1.77]; 8) CYP75A; flavonoid 3&#x2032;,5&#x2032;-hydroxylase [EC:1.14.13.88]; 9) LAR; leucoanthocyanidin reductase [EC:1.17.1.3]; 10) E2.3.1.133; shikimate O-hydroxycinnamoyltransferase [EC:2.3.1.133]; 11) CYP98A; coumaroylquinate (coumaroylshikimate) 3&#x2032;-monooxygenase [EC:1.14.13.36]; 12) E2.1.1.104; caffeoyl-CoA O-methyltransferase [EC:2.1.1.104]; 13) CYP73A; trans-cinnamate 4-monooxygenase [EC:1.14.13.11]; 14) E1.14.11.19; leucoanthocyanidin dioxygenase [EC:1.14.11.19].</p>
</caption>
<graphic xlink:href="fgene-13-786825-g005.tif"/>
</fig>
</sec>
</sec>
<sec id="s3-7">
<title>Comparative Genome Analysis of <italic>P. granatum</italic> and Other Eudicot Species</title>
<p>The full gene sets of pomegranate (<italic>Punica granatum</italic>), apple (<italic>Malus domestica</italic>), <italic>Arabidopsis</italic> (<italic>Arabidopsis thaliana</italic>), grape (<italic>Vitis vinifera</italic>), and <italic>Eucalyptus</italic> (<italic>Eucalyptus grandis</italic>) were analyzed using a gene family cluster analysis and is depicted in <xref ref-type="fig" rid="F6">Figure 6</xref>. The pomegranate genome has 30,803 genes organized into 15,612 gene clusters, 10,435 of which are shared by all five species, advocating their conservation in the lineage after speciation. <italic>P. granatum</italic> shared more gene family clusters (15,316) with <italic>E. grandis</italic> than any of the other three species. Moreover, 1,681 clusters were unique to <italic>P. granatum</italic>. The gene clusters within many genes or in-paralog clusters are most likely the sources of these clusters. The presence of in-paralog clusters suggests that certain gene families in <italic>P. granatum</italic> may have undergone lineage-specific gene expansion. According to the annotation of these clusters, some of these lineage-specific clusters may be involved in key biological processes such as cellular processes, genetic information processing, metabolism, biosynthesis of secondary metabolites, and environmental information processing. The biosynthesis of indole alkaloid, phenylpropanoid, flavonoid, anthocyanin, isoflavonoid, stilbenoid, gingerol, isoquinoline alkaloids, tropane, piperidine and pyridine alkaloid, and glucosinolate genes are identified. However, the present study focused only on phenylpropanoid, flavonoid biosynthesis due to the therapeutic benefits of these secondary metabolites.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Venn diagram of shared orthologous gene families in <italic>Punica granatum</italic>, <italic>Eucalyptus grandis</italic>, <italic>Malus domestica</italic>, <italic>Vitis vinifera</italic>, and <italic>Arabidopsis thaliana</italic>. The gene family number is listed in each component.</p>
</caption>
<graphic xlink:href="fgene-13-786825-g006.tif"/>
</fig>
</sec>
<sec id="s3-8">
<title>Phylogenetic Tree Based on Average Nucleotide Identity Analysis</title>
<p>A hybrid draft genome of <italic>P. granatum</italic> was built using MaSURCA and WENGAN. We compared it to genome sequences of <italic>P. persica</italic> (cultivar_Lovell), <italic>P. granatum</italic> (cultivar Bhagwa, Dabenzi, Taishanhong, Tunisia 2019, and strain AG 2017), <italic>S. alba</italic>, <italic>V. vinifera</italic> (cultivar PN40024), <italic>S. caseolaris</italic>, <italic>E. grandis</italic> (isolate ANBG69807), <italic>C. maculata</italic> (isolate sf003), <italic>A. floribunda</italic> (isolate sf002), and <italic>C. citriodora</italic> (subsp. variegata) The <italic>P. granatum</italic> cultivar Bhagwa was closely related to other <italic>Punica</italic> varieties such as Taishanhong and Tunisia 2019 (ANIm &#x3e;99%). <italic>P. granatum</italic> cultivar Dabenzi, draft assembly, and strain AG2017 also share the same clad and exhibit 99% genetic identity according to average nucleotide analysis using ANI tool PyANI (ANIm) (<xref ref-type="fig" rid="F7">Figure 7</xref>). ANIm genome comparison confirmed that <italic>Eucalyptus grandis</italic> and <italic>Vitis vinifera</italic> are distant relatives of <italic>P. granatum</italic> cultivar Bhagwa falling into separate clad, splitting from other clade members.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Heatmap of ANIm percentage identity: species-level assignments and isolate identifiers as indicated at source given as row and column labels. Cells in the heatmap corresponding to 95% ANIm sequence identity are colored red. Blue cells correspond to ANIm comparisons indicating that the corresponding organisms do not belong to the same species. Color intensity fades as the comparisons approach 95% ANIm sequence identity. Color bars above and to the left of the heatmap correspond to source species-level assignments for each isolate in the analysis. Hierarchical clustering of the analysis results in two dimensions is represented by dendrograms, constructed by simple linkage of ANIm percentage identities.</p>
</caption>
<graphic xlink:href="fgene-13-786825-g007.tif"/>
</fig>
</sec>
<sec id="s3-9">
<title>Identification and Analysis of Variants in <italic>P. granatum G</italic>enome</title>
<p>Variant analysis using the representative genome of <italic>P. granatum</italic> (accession number: GCF_007655135.1_ASM765513v2) from NCBI revealed that the variants are categorized as 359,607 (80.35%) SNPs, 39,857 (8.90%) deletions, and 48,074 (10.74%) insertions. Considering the variants categorized based on the impact, 1,008,520 (95.146%) were classified as modifiers, 28,258 (2.666%) as moderate, and 20,275 (1.913%) as low impact variants. While distributions of the variants based on functional class revealed 27,287 (61.696%) as missense mutations, 16,585 (37.499%) as silent, and 356 (0.805%) as nonsense mutations, the counts based on the effects showed 358,049 (33.668%) in the intergenic region, 100,729 (9.472%) as intron variants, 27,146 (2.553%) missense variants, 16,571 (1.558%) synonymous variants, and 13,546 (1.274%) 3&#x2032;-UTR variants. The variants were classified based on their putative effect on annotated genes, and gene ontology analysis was carried out to acquire a deeper understanding of the functions of 10 genes affected by large-effect variants. These 10 protein sequences were classified into 18 functional classes under three core categories: biological processes (BP), cellular components (CC), and molecular functions (MF), based on homology (<xref ref-type="sec" rid="s11">Supplementary Table S4</xref>).</p>
</sec>
</sec>
<sec sec-type="discussion" id="s4">
<title>Discussion</title>
<p>
<italic>P. granatum</italic>, the crown jewel of the fruit world, is highly nutritious, and produces a phenomenal amount of phytopharmaceuticals and nutraceuticals. The whole draft genome described here in this study using a hybrid approach will serve as a helpful resource in identifying genes and determining their functions. In any genome sequencing project, the foremost requirement is genome size estimation, as it aids in planning the genomic library construction and deciding the amount of raw sequence data to be collected. The nuclear genome size of a plant species is also one of the seminal characteristics, which provides a basic understanding of its cytogenetic features, taxonomic location, and evolution. It is also used to validate the completeness of whole-genome assemblies (<xref ref-type="bibr" rid="B1">Al-Qurainy et al., 2021</xref>). However, primary cytogenetic data is unavailable for <italic>P. granatum</italic> &#x201c;Bhagwa.&#x201d; Therefore, first, we estimated the genome size of <italic>P. granatum</italic> &#x201c;Bhagwa&#x201d; using the flow cytometry method, which was determined to be &#x223c;352.38&#xa0;Mb, &#x223c;76&#xa0;Mb larger than the estimated size of 276.6&#xa0;Mb predicted by GenomeScope, based on Illumina paired-end data.</p>
<p>The whole-genome <italic>de novo</italic> assembly is a critical step in genome research. It is decisive in drawing further genetic resources such as gene annotations, repeat element predictions, and pathway predictions. The datasets obtained using two next-generation sequencing (NGS) techniques and subsequent assembly using several assemblers using various algorithms revealed a difference in genome size estimation using K-mer and flow cytometric analysis. This invoked the need for determining a more reliable and capable assembler for the obtained pomegranate plant genomic data using comparative sequence analysis. The assembled genome size of the <italic>P. granatum</italic> &#x201c;Bhagwa&#x201d; (361.76&#xa0;Mb) is slightly larger than that of the Chinese varieties of pomegranate &#x201c;Dabenzi&#x201d; (328.38&#xa0;Mb) (<xref ref-type="bibr" rid="B38">Qin et al., 2017</xref>), &#x201c;Taishanhong&#x201d; (274&#xa0;Mb) (<xref ref-type="bibr" rid="B50">Yuan et al., 2018</xref>), and &#x201c;Tunisia&#x201d; (320.31&#xa0;Mb) (<xref ref-type="bibr" rid="B28">Luo et al., 2020</xref>). The draft genome completeness is comparable to all other assemblies of <italic>P. granatum</italic> available at NCBI.</p>
<p>In a genome project, genome assembly is succeeded by an ensemble of gene annotations to gain insights into the plant&#x2019;s taxonomy, development, evolution, and functional and metabolic potential (<xref ref-type="bibr" rid="B21">Josep et al., 2019</xref>). TEs exhibit high diversity in structure and modes of transposition and play a vital role in genome evolution, gene regulation, and epigenetics (<xref ref-type="bibr" rid="B39">Sahebi et al., 2018</xref>). The TEs in the <italic>P. granatum</italic> genome occupy only a minority (26.68%) of the genome. This landscape is low compared to other pomegranate cultivars, notably with a burst of TE amplification evident in China cultivars (<xref ref-type="sec" rid="s11">Supplementary Table S5</xref>). This minor quantity reported could be due to the use of hybrid sequencing technology, including long-read technology in the present study. The long-read technology can overcome the low resolution of reconstructing repetitive regions. In general, the quantity of LTR-REs mostly correlates with the genome size of the plant species (<xref ref-type="bibr" rid="B51">Zedek et al., 2010</xref>). The small genome size of the angiosperm <italic>P. granatum</italic> &#x201c;Bhagwa&#x201d; correlates well with the proportion of TEs. Retroelements are the predominant elements other than DNA transposons (<xref ref-type="bibr" rid="B5">Bourque et al., 2018</xref>). In <italic>P. granatum</italic>, the <italic>Copia</italic> is more abundant than <italic>Gypsy</italic>, with the LTR Res elements present in large numbers. Much of the adaptation and growth of <italic>P. granatum</italic> in various agro-climatic conditions can be attributed to TEs domesticated by the pomegranate genome (<xref ref-type="bibr" rid="B29">Makarevitch et al., 2015</xref>). On the other side, the TEs could also contribute to variation in the genome size of various cultivars (<xref ref-type="bibr" rid="B40">SanMiguel et al., 1998</xref>). Our results reveal lesser TEs than in other cultivars such as Thaishanhong (<xref ref-type="bibr" rid="B50">Yuan et al., 2018</xref>) and Dabenzi (<xref ref-type="bibr" rid="B28">Luo et al., 2020</xref>). This observation indicates that different genomes of various cultivars of pomegranate show unique TEs expansion patterns due to other evolutionary processes.</p>
<p>The characterization of the SSRs in <italic>P. granatum</italic> revealed that P1 repeats were the most predominant repeats. SSR frequency decreased with an increase in repeat units in the <italic>P. granatum</italic> genome, which is comparable to that documented in monocots (<italic>Brachypodium</italic>, <italic>Sorghum</italic>, and rice) and dicots (<italic>Arabidopsis</italic>, <italic>Medicago</italic>, and <italic>Populus</italic>) (<xref ref-type="bibr" rid="B42">Sonah et al., 2011</xref>). Genome-wide analysis of SSRs is expected to provide insights into quantitative trait loci (QTL) based selection, plant breeding, genetic linkage mapping, population, and evolutionary genetics of <italic>P. granatum.</italic> The relative abundance of SSRs is relatively high with 369.93 SSRs/Mb and in trend with microsatellite frequency, which is higher in small genomes and is lower in large genomes. This higher density can be expected because of the mutational effects of replication slippage (<xref ref-type="bibr" rid="B34">Morgante et al., 2002</xref>).</p>
<p>The deluge of data from modern genomics technologies empowered research on the biosynthesis and regulation of diverse plant secondary metabolites (<xref ref-type="bibr" rid="B24">Kim and Buell, 2015</xref>). Gene annotations resulted in the prediction of a total of 30,803 proteins. Examination of the genome sequence of <italic>P. granatum</italic> also enabled the identification of genomic signatures of secondary metabolism genes of phenylpropanoid (17 genes) and flavonoid (14 genes) pathways. <italic>P. granatum</italic> is a good source of p-coumaric acid and is one of the very important phytopharmaceuticals with anti-breast cancer activity, as reported in one of the earlier studies (<xref ref-type="bibr" rid="B44">Usha et al., 2021</xref>). Notably, it is also a precursor molecule for flavonoid biosynthesis (<xref ref-type="bibr" rid="B17">Ferreyra et al., 2012</xref>). The <italic>P. granatum</italic> produces a high amount of polyphenols and flavonoids (<xref ref-type="bibr" rid="B32">Middha et al., 2013a</xref>), which highly correlates with the existence of a more significant amount of phenylpropanoid and flavonoid pathway genes in <italic>P. granatum</italic> &#x201c;Bhagwa.&#x201d; The enzymes that catalyze the synthesis of major flavonoids, cyanidin, epicatechin, kaempferol, luteolin, naringin, pelargonidin, and quercetin, were identified. These flavonoids exhibit antioxidant, antineoplastic, anti-inflammatory, antiviral, and antibacterial activities (<xref ref-type="bibr" rid="B33">Middha et al., 2013b</xref>). The genome sequence of <italic>P. granatum</italic> &#x201c;Dabenzi&#x201d; provides information on candidate genes for punicalagin biosynthesis, which mainly catalyzes gallic acid synthesis from shikimic acid and syringate, not from quinic acid (<xref ref-type="bibr" rid="B38">Qin et al., 2017</xref>). These results coincide with our findings that phenylpropanoids are synthesized from shikimic acid in <italic>P. granatum</italic> &#x201c;Bhagwa&#x201d; species. Consistent with prior reports, the recent common ancestry between <italic>E</italic>. <italic>grandis</italic> and <italic>P</italic>. <italic>granatum</italic> was observed through comparative genome analysis studies of clustered single-copy gene orthologs related to secondary metabolism (<xref ref-type="bibr" rid="B38">Qin et al., 2017</xref>). The current study illustrates how hybrid sequencing technology can resolve complex TE and SSRs. These regions encompass essential genomic information critical for the adaptation and evolution to the environment, thus assisting in developing the crops by genetic breeding methodologies. The draft genome and its annotations of <italic>P. granatum</italic> &#x201c;Bhagwa&#x201d; will accelerate crop improvement by selecting desirable genes with enhanced agronomic traits, including nutrient richness, high yield, biotic and abiotic stress tolerance, and resistance against pathogens, and high yield of secondary metabolites with pharmacological properties.</p>
</sec>
<sec sec-type="conclusion" id="s5">
<title>Conclusion</title>
<p>The assembled first draft genome of <italic>P. granatum</italic> &#x201c;Bhagwa&#x201d; in this study can be a valuable resource and reference to understanding this commercially important fruit crop&#x2019;s taxonomy, evolution, and biological architecture. The study also provides a novel experimental (hybrids sequencing technology) and computational approach of using multiple assemblers for dealing with the difference between the flow cytometric and k-mer genome size estimations. The vast data and information obtained from the draft genome of <italic>P. granatum</italic> will also improve the crop identification and understanding of the biosynthesis of phytopharmaceuticals. Unraveling the high-quality complete genome and transcriptome of the <italic>P. granatum</italic> will further facilitate future research into other areas of research on this plant, such as aspects of its environmental stress tolerance, acclimatization, evolution, biosynthesis of other pharmacologically important secondary metabolites, crop improvement, and resistance to pathogens.</p>
</sec>
</body>
<back>
<sec id="s6">
<title>Data Availability Statement</title>
<p>All the raw data used in this study was submitted to the NCBI SRA data repository (BioSample: SAMN07645014; Sample name: Pomegranate (<italic>Punica granatum</italic>); SRA: SRS2645763) (<ext-link ext-link-type="uri" xlink:href="https://www.ncbi.nlm.nih.gov/bioproject/PRJNA407279">https://www.ncbi.nlm.nih.gov/bioproject/PRJNA407279</ext-link>) under the accession number. All secondary data used in this study are available in the <xref ref-type="sec" rid="s11">Supplementary Materials</xref> provided.</p>
</sec>
<sec id="s7">
<title>Author Contributions</title>
<p>SM, DS and KS conceptualized the whole study. TU performed DNA isolations, genomic library constructions, genome assembly, and genome annotations and drafted the manuscript. DS was involved in sequencing genome libraries using Illumina and QC. VK prepared and sequenced genomic libraries for Oxford Nanopore sequencing. BR assisted in assembling the genome and phylogenetic analysis. AG helped in the experimental procedure and partially drafted the manuscript. AJD and AS helped in the genome assembly and annotations. DB supervised the computational analysis and thoroughly revised the manuscript. SM and DS arranged the funds, supervised the whole study, and edited the final version of the manuscript. PM bred the pure lines of <italic>P. granatum</italic> &#x201c;Bhagwa,&#x201d; supervised all the experimental parts of the work, and wrote part of the manuscript. All the authors analyzed and discussed the results and approved this submitted version.</p>
</sec>
<sec id="s8">
<title>Funding</title>
<p>We are thankful to BiSEP, Govt of Karnataka, for partially supporting this study and the DBT-BIF centre at MLACW for providing a computational facility.</p>
</sec>
<sec sec-type="COI-statement" id="s9">
<title>Conflict of Interest</title>
<p>AD was employed by Molsys Pvt. Ltd. DP was employed by Protein Design Private Limited. VK was employed by Genotypic Technology Pvt Limited. AS was employed by Basesolve Informatics Pvt Ltd.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x2019;s Note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations or those of the publisher, the editors, and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ack>
<p>We acknowledge Cytometry Solutions Private Limited for the flow cytometry facility and their technical support. DBT-BIF computational facility and DST-FIST facility at MLACW were used to carry out the research. The authors also acknowledge the suggestions provided by V. R. Devraj, Professor, Bangalore City University, India, and C. S. Karigar, Professor, Bangalore University, India.</p>
</ack>
<sec id="s11">
<title>Supplementary Material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fgene.2022.786825/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fgene.2022.786825/full&#x23;supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Table3.xlsx" id="SM1" mimetype="application/xlsx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table1.docx" id="SM2" mimetype="application/docx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="DataSheet1.doc" id="SM3" mimetype="application/doc" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image1.jpeg" id="SM4" mimetype="application/jpeg" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image2.jpeg" id="SM5" mimetype="application/jpeg" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table5.docx" id="SM6" mimetype="application/docx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table4.xlsx" id="SM7" mimetype="application/xlsx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table2.doc" id="SM8" mimetype="application/doc" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<fn-group>
<fn id="fn1">
<label>1</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://clinicaltrials.gov/ct2/results?term=pomegranateandSearch=Applyandrecrs=eandage_v=andgndr=andtype=andrslt=">https://clinicaltrials.gov/ct2/results?term&#x003D;pomegranateandSearch&#x003D;Applyandrecrs&#x003D;eandage_v&#x003D;andgndr&#x003D;andtype&#x003D;andrslt&#x003D;</ext-link>.</p>
</fn>
<fn id="fn2">
<label>2</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://github.com/rrwick/Filtlong)(parameters%20--min_length%20500%20--keep_percent%2090">https://github.com/rrwick/Filtlong)(parameters --min_length 500 --keep_percent 90</ext-link>.</p>
</fn>
<fn id="fn3">
<label>3</label>
<p>
<ext-link ext-link-type="uri" xlink:href="http://www.genome.umd.edu/masurca.html">http://www.genome.umd.edu/masurca.html</ext-link>.</p>
</fn>
<fn id="fn4">
<label>4</label>
<p>
<ext-link ext-link-type="uri" xlink:href="http://www.repeatmasker.org/RMBlast.html">http://www.repeatmasker.org/RMBlast.html</ext-link>.</p>
</fn>
<fn id="fn5">
<label>5</label>
<p>
<ext-link ext-link-type="uri" xlink:href="http://pgrc.ipk-gatersleben.de/misa/">sleben.de/misa/</ext-link>.</p>
</fn>
<fn id="fn6">
<label>6</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://github.com/mnenno/PySSRstat">https://github.com/mnenno/PySSRstat</ext-link>.</p>
</fn>
<fn id="fn7">
<label>7</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://github.com/tseemann/barrnap">https://github.com/tseemann/barrnap</ext-link>.</p>
</fn>
<fn id="fn8">
<label>8</label>
<p>
<ext-link ext-link-type="uri" xlink:href="http://ccb.jhu.edu/software/stringtie/gff.shtml%20\l%20gffread">http://ccb.jhu.edu/software/stringtie/gff.shtml&#x23;gffread</ext-link>.</p>
</fn>
<fn id="fn9">
<label>9</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://www.genome.jp/kegg/kaas/">https://www.genome.jp/kegg/kaas/</ext-link>.</p>
</fn>
<fn id="fn10">
<label>10</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://gatk.broadinstitute.org/hc/en-us">https://gatk.broadinstitute.org/hc/en-us</ext-link>.</p>
</fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Al-Qurainy</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Gaafar</surname>
<given-names>A.-R. Z.</given-names>
</name>
<name>
<surname>Khan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Nadeem</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Alshameri</surname>
<given-names>A. M.</given-names>
</name>
<name>
<surname>Tarroum</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Estimation of Genome Size in the Endemic Species <italic>Reseda Pentagyna</italic> and the Locally Rare Species <italic>Reseda Lutea</italic> Using Comparative Analyses of Flow Cytometry and K-Mer Approaches</article-title>. <source>Plants</source> <volume>10</volume> (<issue>7</issue>), <fpage>1362</fpage>. <pub-id pub-id-type="doi">10.3390/plants10071362</pub-id> </citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Alonge</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Soyk</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ramakrishnan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Goodwin</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Sedlazeck</surname>
<given-names>F. J.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>RaGOO: Fast and Accurate Reference-Guided Scaffolding of Draft Genomes</article-title>. <source>Genome Biol.</source> <volume>20</volume> (<issue>1</issue>), <fpage>224</fpage>&#x2013;<lpage>317</lpage>. <pub-id pub-id-type="doi">10.1186/s13059-019-1829-6</pub-id> </citation>
</ref>
<ref id="B3">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Andrews</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>FastQC: a Quality Control Tool for High Throughput Sequence Data</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="http://www.bioinformatics.babraham.ac.uk/projects/fastqc">http://www.bioinformatics.babraham.ac.uk/projects/fastqc</ext-link>
</comment>. </citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Balasubramani</surname>
<given-names>S. P.</given-names>
</name>
<name>
<surname>Mohan</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Chatterjee</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Patnaik</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Kukkupuni</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Nongthomba</surname>
<given-names>U.</given-names>
</name>
<etal/>
</person-group> (<year>2014</year>). <article-title>Pomegranate Juice Enhances Healthy Lifespan in <italic>Drosophila melanogaster</italic>: an Exploratory Study</article-title>. <source>Front. Public Health</source> <volume>2</volume>, <fpage>245</fpage>. <pub-id pub-id-type="doi">10.3389/fpubh.2014.00245</pub-id> </citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bourque</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Burns</surname>
<given-names>K. H.</given-names>
</name>
<name>
<surname>Gehring</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Gorbunova</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Seluanov</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Hammell</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>Ten Things You Should Know about Transposable Elements</article-title>. <source>Genome Biol.</source> <volume>19</volume>, <fpage>199</fpage>. <pub-id pub-id-type="doi">10.1186/s13059-018-1577-z</pub-id> </citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Br&#x16f;na</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Hoff</surname>
<given-names>K. J.</given-names>
</name>
<name>
<surname>Lomsadze</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Stanke</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Borodovsky</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>BRAKER2: Automatic Eukaryotic Genome Annotation with GeneMark-Ep&#x2b; and AUGUSTUS Supported by a Protein Database</article-title>. <source>NAR Genom. Bioinform.</source> <volume>3</volume> (<issue>1</issue>), <fpage>lqaa108</fpage>. <pub-id pub-id-type="doi">10.1093/nargab/lqaa108</pub-id> </citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Buchfink</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Huson</surname>
<given-names>D. H.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Fast and Sensitive Protein Alignment Using DIAMOND</article-title>. <source>Nat. Methods</source> <volume>12</volume>, <fpage>59</fpage>&#x2013;<lpage>60</lpage>. <pub-id pub-id-type="doi">10.1038/nmeth.3176</pub-id> </citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chan</surname>
<given-names>P. P.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>B. Y.</given-names>
</name>
<name>
<surname>AllysiaMak</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>ToddLowe</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>tRNAscan-SE 2.0: Improved Detection and Functional Classification of Transfer RNA Genes</article-title>. <source>Nucleic Acids Res.</source> <volume>49</volume> (<issue>16</issue>), <fpage>9077</fpage>&#x2013;<lpage>9096</lpage>. <pub-id pub-id-type="doi">10.1093/nar/gkab688</pub-id> </citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chandra</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Jadhav</surname>
<given-names>V. T.</given-names>
</name>
<name>
<surname>Sharma</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Global Scenario of Pomegranate (<italic>Punica Granatum</italic> L.) Culture with Special Reference to India</article-title>. <source>Fruit Veg. Cereal Sci. Biotechnol. Spec. Issue</source> <volume>2</volume>, <fpage>7</fpage>&#x2013;<lpage>18</lpage>. </citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Gu</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Fastp: an Ultra-fast All-In-One FASTQ Preprocessor</article-title>. <source>Bioinformatics</source> <volume>34</volume> (<issue>17</issue>), <fpage>i884</fpage>&#x2013;<lpage>i890</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/bty560</pub-id> </citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Nie</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>S. Q.</given-names>
</name>
<name>
<surname>Zheng</surname>
<given-names>Y. F.</given-names>
</name>
<name>
<surname>Dai</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Bray</surname>
<given-names>T.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Efficient Assembly of Nanopore Reads via Highly Accurate and Intact Error Correction</article-title>. <source>Nat. Commun.</source> <volume>12</volume> (<issue>60</issue>). <pub-id pub-id-type="doi">10.1038/s41467-020-20236-7</pub-id> </citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cingolani</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Platts</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>L. L.</given-names>
</name>
<name>
<surname>Coon</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>L.</given-names>
</name>
<etal/>
</person-group> (<year>2012</year>). <article-title>A Program for Annotating and Predicting the Effects of Single Nucleotide Polymorphisms, SnpEff</article-title>. <source>Fly</source> <volume>6</volume>, <fpage>80</fpage>&#x2013;<lpage>92</lpage>. <pub-id pub-id-type="doi">10.4161/fly.19695</pub-id> </citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Conesa</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>G&#xf6;tz</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Garc&#xed;a-G&#xf3;mez</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Terol</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Tal&#xf3;n</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Robles</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>Blast2GO: a Universal Tool for Annotation, Visualization and Analysis in Functional Genomics Research</article-title>. <source>Bioinformatics</source> <volume>21</volume> (<issue>18</issue>), <fpage>3674</fpage>&#x2013;<lpage>3676</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/bti610</pub-id> </citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>DePristo</surname>
<given-names>M. A.</given-names>
</name>
<name>
<surname>Banks</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Poplin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Garimella</surname>
<given-names>K. V.</given-names>
</name>
<name>
<surname>Maguire</surname>
<given-names>J. R.</given-names>
</name>
<name>
<surname>Hartl</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2011</year>). <article-title>A Framework for Variation Discovery and Genotyping Using Next-Generation DNA Sequencing Data</article-title>. <source>Nat. Genet.</source> <volume>43</volume>, <fpage>491</fpage>&#x2013;<lpage>498</lpage>. <pub-id pub-id-type="doi">10.1038/ng.806</pub-id> </citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Di Genova</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Buena-Atienza</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Ossowski</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Sagot</surname>
<given-names>M.-F.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Efficient Hybrid De Novo Assembly of Human Genomes with WENGAN</article-title>. <source>Nat. Biotechnol.</source> <volume>39</volume>, <fpage>422</fpage>&#x2013;<lpage>430</lpage>. <pub-id pub-id-type="doi">10.1038/s41587-020-00747-w</pub-id> </citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fadavi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Barzegar</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Azizi</surname>
<given-names>M. H.</given-names>
</name>
<name>
<surname>Bayat</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>Note. Physicochemical Composition of Ten Pomegranate Cultivars (Punica Granatum L.) Grown in Iran</article-title>. <source>Food Sci. Technol. Int.</source> <volume>11</volume>, <fpage>113</fpage>&#x2013;<lpage>119</lpage>. <pub-id pub-id-type="doi">10.1177/1082013205052765</pub-id> </citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ferreyra</surname>
<given-names>M. L. F.</given-names>
</name>
<name>
<surname>Rius</surname>
<given-names>S. P.</given-names>
</name>
<name>
<surname>Casati</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Flavonoids: Biosynthesis, Biological Functions, and Biotechnological Applications</article-title>. <source>Front. Plant Sci.</source> <volume>3</volume>, <fpage>222</fpage>. <pub-id pub-id-type="doi">10.3389/fpls.2012.00222</pub-id> </citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Flynn</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Hubley</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Goubert</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Rosen</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Clark</surname>
<given-names>A. G.</given-names>
</name>
<name>
<surname>Feschotte</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>RepeatModeler2 for Automated Genomic Discovery of Transposable Element Families</article-title>. <source>Proc. Natl. Acad. Sci. U.S.A.</source> <volume>117</volume> (<issue>17</issue>), <fpage>9451</fpage>&#x2013;<lpage>9457</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.1921046117</pub-id> </citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fukasawa</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Ermini</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Carty</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Cheung</surname>
<given-names>M.-S.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Long QC: A Quality Control Tool for Third Generation Sequencing Long Read Data [published Correction Appears in G3 (Bethesda)</article-title>. <source>G3: Genes, Genomes, Genet. (Bethesda)</source> <volume>10</volume> (<issue>4</issue>), <fpage>1193</fpage>&#x2013;<lpage>1196</lpage>. <pub-id pub-id-type="doi">10.1534/g3.119.400864</pub-id> </citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Galbraith</surname>
<given-names>D. W.</given-names>
</name>
<name>
<surname>Harkins</surname>
<given-names>K. R.</given-names>
</name>
<name>
<surname>Maddox</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Ayres</surname>
<given-names>N. M.</given-names>
</name>
<name>
<surname>Sharma</surname>
<given-names>D. P.</given-names>
</name>
<name>
<surname>Firoozabady</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>1983</year>). <article-title>Rapid Flow Cytometric Analysis of the Cell Cycle in Intact Plant Tissues</article-title>. <source>Science</source> <volume>220</volume>, <fpage>1049</fpage>&#x2013;<lpage>1051</lpage>. <pub-id pub-id-type="doi">10.1126/science.220.4601.1049</pub-id> </citation>
</ref>
<ref id="B21">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Josep</surname>
<given-names>F. A.</given-names>
</name>
<name>
<surname>Castellano</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>Genome Annotation</article-title>,&#x201d; in <source>Encyclopedia of Bioinformatics and Computational Biology</source>. Editors <person-group person-group-type="editor">
<name>
<surname>Ranganathan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Gribskov</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Nakai</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Sch&#xf6;nbach</surname>
<given-names>C.</given-names>
</name>
</person-group> (<publisher-loc>Cambridge</publisher-loc>: <publisher-name>Elsevier Academic Press</publisher-name>), <fpage>195</fpage>&#x2013;<lpage>209</lpage>. </citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jurenka</surname>
<given-names>J. S.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Therapeutic Applications of Pomegranate (<italic>Punica Granatum</italic> L.): a Review</article-title>. <source>Altern. Med. Rev.</source> <volume>13</volume>, <fpage>128</fpage>&#x2013;<lpage>144</lpage>. </citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kandylis</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Kokkinomagoulos</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Food Applications and Potential Health Benefits of Pomegranate and its Derivatives</article-title>. <source>Foods</source> <volume>9</volume>, <fpage>122</fpage>. <pub-id pub-id-type="doi">10.3390/foods9020122</pub-id> </citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kim</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Buell</surname>
<given-names>C. R.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>A Revolution in Plant Metabolism: Genome-Enabled Pathway Discovery</article-title>. <source>Plant Physiol.</source> <volume>169</volume>, <fpage>1532</fpage>&#x2013;<lpage>1539</lpage>. <pub-id pub-id-type="doi">10.1104/pp.15.00976</pub-id> </citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lazare</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lyu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yermiyahu</surname>
<given-names>U.</given-names>
</name>
<name>
<surname>Heler</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kalyan</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Dag</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>The Effect of Macronutrient Availability on Pomegranate Reproductive Development</article-title>. <source>Plants</source> <volume>9</volume>, <fpage>963</fpage>. <pub-id pub-id-type="doi">10.3390/plants9080963</pub-id> </citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Handsaker</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Wysoker</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Fennell</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Ruan</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Homer</surname>
<given-names>N.</given-names>
</name>
<etal/>
</person-group> (<year>2009</year>). <article-title>The Sequence Alignment/map Format and SAMtools</article-title>. <source>Bioinformatics</source> <volume>25</volume>, <fpage>2078</fpage>&#x2013;<lpage>2079</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btp352</pub-id> </citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liao</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Gao</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A Sensitive Repeat Identification Framework Based on Short and Long Reads</article-title>. <source>Nucleic Acids Res.</source> <volume>49</volume> (<issue>17</issue>), <fpage>e100</fpage>. <pub-id pub-id-type="doi">10.1093/nar/gkab563</pub-id> </citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Luo</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>The Pomegranate (<italic>Punica Granatum</italic> L.) Draft Genome Dissects Genetic Divergence between Soft- and Hard-Seeded Cultivars</article-title>. <source>Plant Biotechnol. J.</source> <volume>18</volume>, <fpage>955</fpage>&#x2013;<lpage>968</lpage>. <pub-id pub-id-type="doi">10.1111/pbi.13260</pub-id> </citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Makarevitch</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Waters</surname>
<given-names>A. J.</given-names>
</name>
<name>
<surname>West</surname>
<given-names>P. T.</given-names>
</name>
<name>
<surname>Stitzer</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Hirsch</surname>
<given-names>C. N.</given-names>
</name>
<name>
<surname>Ross-Ibarra</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>20152015</year>). <article-title>Transposable Elements Contribute to Activation of maize Genes in Response to Abiotic Stress</article-title>. <source>Plos Genet.</source> <volume>11</volume>, <fpage>e1004915</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pgen.1004915</pub-id> </citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Martin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Hackl</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Hattab</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Fischer</surname>
<given-names>M. G.</given-names>
</name>
<name>
<surname>Heider</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>MOSGA: Modular Open-Source Genome Annotator</article-title>. <source>Bioinformatics</source> <volume>36</volume> (<issue>22-23</issue>), <fpage>5514</fpage>&#x2013;<lpage>5515</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btaa1003</pub-id> </citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Medjakovic</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Jungbauer</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>Pomegranate: a Fruit that Ameliorates Metabolic Syndrome</article-title>. <source>Food Funct.</source> <volume>4</volume>, <fpage>19</fpage>&#x2013;<lpage>39</lpage>. <pub-id pub-id-type="doi">10.1039/c2fo30034f</pub-id> </citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Middha</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Usha</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Pande</surname>
<given-names>V.</given-names>
</name>
</person-group> (<year>2013a</year>). <article-title>A Review on Antihyperglycemic and Antihepatoprotective Activity of Eco-Friendly <italic>Punica Granatum</italic> Peel Waste. Evid. Based Complement</article-title>. <source>Alternat. Med.</source> <volume>2013</volume>, <fpage>656172</fpage>. <pub-id pub-id-type="doi">10.1155/2013/656172</pub-id> </citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Middha</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Usha</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Pande</surname>
<given-names>V.</given-names>
</name>
</person-group> (<year>2013b</year>). <article-title>HPLC Evaluation of Phenolic Profile, Nutritive Content and Antioxidant Capacity of Extracts Obtained from <italic>Punica Granatum</italic> Fruit Peel</article-title>. <source>Adv. Pharmacol. Sci.</source> <volume>2013</volume>, <fpage>296236</fpage>. <pub-id pub-id-type="doi">10.1155/2013/296236</pub-id> </citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Morgante</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Hanafey</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Powell</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>2002</year>). <article-title>Microsatellites Are Preferentially Associated with Nonrepetitive DNA in Plant Genomes</article-title>. <source>Nat. Genet.</source> <volume>30</volume>, <fpage>194</fpage>&#x2013;<lpage>200</lpage>. <pub-id pub-id-type="doi">10.1038/ng822</pub-id> </citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ou</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Su</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Liao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chougule</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Agda</surname>
<given-names>J. R.</given-names>
</name>
<name>
<surname>Hellinga</surname>
<given-names>A. J.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Benchmarking Transposable Element Annotation Methods for Creation of a Streamlined, Comprehensive Pipeline</article-title>. <source>Genome Biol.</source> <volume>20</volume>, <fpage>275</fpage>. <pub-id pub-id-type="doi">10.1186/s13059-019-1905-y</pub-id> </citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pramesh</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Prasannakumar</surname>
<given-names>M. K.</given-names>
</name>
<name>
<surname>Muniraju</surname>
<given-names>K. M.</given-names>
</name>
<name>
<surname>Mahesh</surname>
<given-names>H. B.</given-names>
</name>
<name>
<surname>Pushpa</surname>
<given-names>H. D.</given-names>
</name>
<name>
<surname>Manjunatha</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Comparative Genomics of rice False Smut Fungi <italic>Ustilaginoidea Virens</italic> Uv-Gvt Strain from India Reveals Genetic Diversity and Phylogenetic Divergence</article-title>. <source>Biotech</source> <volume>10</volume> (<issue>8</issue>), <fpage>342</fpage>. <pub-id pub-id-type="doi">10.1007/s13205-020-02336-9</pub-id> </citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pryszcz</surname>
<given-names>L. P.</given-names>
</name>
<name>
<surname>Gabald&#xf3;n</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Redundans: an Assembly Pipeline for Highly Heterozygous Genomes</article-title>. <source>Nucleic Acids Res.</source> <volume>44</volume> (<issue>12</issue>), <fpage>e113</fpage>. <pub-id pub-id-type="doi">10.1093/nar/gkw294</pub-id> </citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qin</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Ming</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Tang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Guyot</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Kramer</surname>
<given-names>E. M.</given-names>
</name>
<etal/>
</person-group> (<year>2017</year>). <article-title>The Pomegranate (<italic>Punica Granatum</italic> L.) Genome and the Genomics of Punicalagin Biosynthesis</article-title>. <source>Plant J.</source> <volume>91</volume>, <fpage>1108</fpage>&#x2013;<lpage>1128</lpage>. <pub-id pub-id-type="doi">10.1111/tpj.13625</pub-id> </citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sahebi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Hanafi</surname>
<given-names>M. M.</given-names>
</name>
<name>
<surname>Van Wijnen</surname>
<given-names>A. J.</given-names>
</name>
<name>
<surname>Rice</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Rafii</surname>
<given-names>M. Y.</given-names>
</name>
<name>
<surname>Azizi</surname>
<given-names>P.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>Contribution of Transposable Elements in the Plant&#x27;s Genome</article-title>. <source>Gene</source> <volume>665</volume>, <fpage>155</fpage>&#x2013;<lpage>166</lpage>. <pub-id pub-id-type="doi">10.1016/j.gene.2018.04.050</pub-id> </citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>SanMiguel</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Gaut</surname>
<given-names>B. S.</given-names>
</name>
<name>
<surname>Tikhonov</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Nakajima</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Bennetzen</surname>
<given-names>J. L.</given-names>
</name>
</person-group> (<year>1998</year>). <article-title>The Paleontology of Intergene Retrotransposons of maize</article-title>. <source>Nat. Genet.</source> <volume>20</volume>, <fpage>43</fpage>&#x2013;<lpage>45</lpage>. <pub-id pub-id-type="doi">10.1038/1695</pub-id> </citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sim&#xe3;o</surname>
<given-names>F. A.</given-names>
</name>
<name>
<surname>Waterhouse</surname>
<given-names>R. M.</given-names>
</name>
<name>
<surname>Ioannidis</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Kriventseva</surname>
<given-names>E. V.</given-names>
</name>
<name>
<surname>Zdobnov</surname>
<given-names>E. M.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>BUSCO: Assessing Genome Assembly and Annotation Completeness with Single-Copy Orthologs</article-title>. <source>Bioinformatics</source> <volume>31</volume> (<issue>19</issue>), <fpage>3210</fpage>&#x2013;<lpage>3212</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btv351</pub-id> </citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sonah</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Deshmukh</surname>
<given-names>R. K.</given-names>
</name>
<name>
<surname>Sharma</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Singh</surname>
<given-names>V. P.</given-names>
</name>
<name>
<surname>Gupta</surname>
<given-names>D. K.</given-names>
</name>
<name>
<surname>Gacche</surname>
<given-names>R. N.</given-names>
</name>
<etal/>
</person-group> (<year>2011</year>). <article-title>Genome-wide Distribution and Organization of Microsatellites in Plants: an Insight into Marker Development in Brachypodium</article-title>. <source>PLoS One</source> <volume>6</volume>, <fpage>e21298</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0021298</pub-id> </citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sreekumar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Sithul</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Muraleedharan</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Azeez</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Sreeharshan</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Pomegranate Fruit as a Rich Source of Biologically Active Compounds</article-title>. <source>Biomed. Res. Int.</source> <volume>2014</volume>, <fpage>686921</fpage>. <pub-id pub-id-type="doi">10.1155/2014/686921</pub-id> </citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Usha</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Middha</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Shanmugarajan</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Babu</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Goyal</surname>
<given-names>A. K.</given-names>
</name>
<name>
<surname>Yusufoglu</surname>
<given-names>H. S.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Gas Chromatography-Mass Spectrometry Metabolic Profiling, Molecular Simulation and Dynamics of Diverse Phytochemicals of <italic>Punica Granatum</italic> L. Leaves against Estrogen Receptor</article-title>. <source>Front. Biosci. Landmark.</source> <volume>26</volume> (<issue>9</issue>), <fpage>423</fpage>&#x2013;<lpage>441</lpage>. <pub-id pub-id-type="doi">10.52586/4957</pub-id> </citation>
</ref>
<ref id="B45">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Usha</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Middha</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Sidhalinghamurthy</surname>
<given-names>K. R.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>Pomegranate Peel and its Anticancer Activity: A Mechanism-Based Review</article-title>,&#x201d; in <source>Plant-derived Bioactives</source>. Editor <person-group person-group-type="editor">
<name>
<surname>Swamy</surname>
<given-names>M.</given-names>
</name>
</person-group> (<publisher-loc>Singapore</publisher-loc>: <publisher-name>Springer</publisher-name>), <fpage>223</fpage>&#x2013;<lpage>250</lpage>. <pub-id pub-id-type="doi">10.1007/978-981-15-2361-8_10</pub-id> </citation>
</ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vu&#x10d;i&#x107;</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Grabe&#x17e;</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Trchounian</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Arsi&#x107;</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>Composition and Potential Health Benefits of Pomegranate: A Review</article-title>. <source>Curr. Pharm. Des.</source> <volume>25</volume>, <fpage>1817</fpage>&#x2013;<lpage>1827</lpage>. <pub-id pub-id-type="doi">10.52586/4957</pub-id> </citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vurture</surname>
<given-names>G. W.</given-names>
</name>
<name>
<surname>Sedlazeck</surname>
<given-names>F. J.</given-names>
</name>
<name>
<surname>Nattestad</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Underwood</surname>
<given-names>C. J.</given-names>
</name>
<name>
<surname>Fang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Gurtowski</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2017</year>). <article-title>GenomeScope: Fast Reference-free Genome Profiling from Short Reads</article-title>. <source>Bioinformatics</source> <volume>33</volume> (<issue>14</issue>), <fpage>2202</fpage>&#x2013;<lpage>2204</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btx153</pub-id> </citation>
</ref>
<ref id="B48">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>West</surname>
<given-names>P. T.</given-names>
</name>
<name>
<surname>Probst</surname>
<given-names>A. J.</given-names>
</name>
<name>
<surname>Grigoriev</surname>
<given-names>I. V.</given-names>
</name>
<name>
<surname>Thomas</surname>
<given-names>B. C.</given-names>
</name>
<name>
<surname>Banfield</surname>
<given-names>J. F.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Genome-reconstruction for Eukaryotes from Complex Natural Microbial Communities</article-title>. <source>Genome Res.</source> <volume>28</volume> (<issue>4</issue>), <fpage>569</fpage>&#x2013;<lpage>580</lpage>. <pub-id pub-id-type="doi">10.1101/gr.228429.117</pub-id> </citation>
</ref>
<ref id="B49">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yan</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Bombarely</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>DeepTE: a Computational Method for De Novo Classification of Transposons with Convolutional Neural Network</article-title>. <source>Bioinformatics</source> <volume>36</volume> (<issue>15</issue>), <fpage>4269</fpage>&#x2013;<lpage>4275</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btaa519</pub-id> </citation>
</ref>
<ref id="B50">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yuan</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Fang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Fei</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>The Pomegranate (<italic>Punica Granatum</italic> L.) Genome Provides Insights into Fruit Quality and Ovule Developmental Biology</article-title>. <source>Plant Biotechnol. J.</source> <volume>16</volume>, <fpage>1363</fpage>&#x2013;<lpage>1374</lpage>. <pub-id pub-id-type="doi">10.1111/pbi.12875</pub-id> </citation>
</ref>
<ref id="B51">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zedek</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>&#x160;merda</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>&#x160;marda</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Bures</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Correlated Evolution of LTR Retrotransposons and Genome Size in the Genus</article-title>. <source>Eleocharisbmc Plant Biol.</source> <volume>10</volume>, <fpage>265</fpage>. <pub-id pub-id-type="doi">10.1186/1471-2229-10-265</pub-id> </citation>
</ref>
<ref id="B52">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zimin</surname>
<given-names>A. V.</given-names>
</name>
<name>
<surname>Mar&#xe7;ais</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Puiu</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Roberts</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Salzberg</surname>
<given-names>S. L.</given-names>
</name>
<name>
<surname>Yorke</surname>
<given-names>J. A.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>The MaSuRCA Genome Assembler</article-title>. <source>Bioinformatics</source> <volume>29</volume>, <fpage>2669</fpage>&#x2013;<lpage>1677</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btt476</pub-id> </citation>
</ref>
</ref-list>
</back>
</article>