<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Plant Sci.</journal-id>
<journal-title>Frontiers in Plant Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Plant Sci.</abbrev-journal-title>
<issn pub-type="epub">1664-462X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpls.2023.1260393</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Plant Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Differential gene expression provides leads to environmentally regulated soybean seed protein content</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Hooker</surname>
<given-names>Julia C.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/811253"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Smith</surname>
<given-names>Myron</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Zapata</surname>
<given-names>Gerardo</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Charette</surname>
<given-names>Martin</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Luckert</surname>
<given-names>Doris</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Mohr</surname>
<given-names>Ramona M.</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/550445"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Daba</surname>
<given-names>Ketema A.</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/300007"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Warkentin</surname>
<given-names>Thomas D.</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/237630"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Hadinezhad</surname>
<given-names>Mehri</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Barlow</surname>
<given-names>Brent</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Hou</surname>
<given-names>Anfu</given-names>
</name>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1438191"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Lefebvre</surname>
<given-names>Fran&#xe7;ois</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Golshani</surname>
<given-names>Ashkan</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Cober</surname>
<given-names>Elroy R.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2322249"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Samanfar</surname>
<given-names>Bahram</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>*</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1717060"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Ottawa Research and Development Centre, Agriculture and Agri-Food Canada</institution>, <addr-line>Ottawa, ON</addr-line>, <country>Canada</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Department of Biology, Ottawa Institute of Systems Biology, Carleton University</institution>, <addr-line>Ottawa, ON</addr-line>, <country>Canada</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Canadian Centre for Computational Genomics</institution>, <addr-line>Montr&#xe9;al, QC</addr-line>, <country>Canada</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Brandon Research Centre, Agriculture and Agri-Food Canada</institution>, <addr-line>Brandon, MB</addr-line>, <country>Canada</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Crop Development Centre, University of Saskatchewan</institution>, <addr-line>Saskatoon, SK</addr-line>, <country>Canada</country>
</aff>
<aff id="aff6">
<sup>6</sup>
<institution>Morden Research and Development Centre, Agriculture and Agri-Food Canada</institution>, <addr-line>Morden, MB</addr-line>, <country>Canada</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>Edited by: Kyung Do Kim, Myongji University, Republic of Korea</p>
</fn>
<fn fn-type="edited-by">
<p>Reviewed by: Huatao Chen, Jiangsu Academy of Agricultural Sciences (JAAS), China; Yong-Qiang Charles An, United States Department of Agriculture (USDA), United States</p>
</fn>
<fn fn-type="corresp" id="fn001">
<p>*Correspondence: Bahram Samanfar, <email xlink:href="mailto:bahram.samanfar@agr.gc.ca">bahram.samanfar@agr.gc.ca</email>
</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>18</day>
<month>09</month>
<year>2023</year>
</pub-date>
<pub-date pub-type="collection">
<year>2023</year>
</pub-date>
<volume>14</volume>
<elocation-id>1260393</elocation-id>
<history>
<date date-type="received">
<day>17</day>
<month>07</month>
<year>2023</year>
</date>
<date date-type="accepted">
<day>23</day>
<month>08</month>
<year>2023</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2023 Hooker, Smith, Zapata, Charette, Luckert, Mohr, Daba, Warkentin, Hadinezhad, Barlow, Hou, Lefebvre, Golshani, Cober and Samanfar</copyright-statement>
<copyright-year>2023</copyright-year>
<copyright-holder>Hooker, Smith, Zapata, Charette, Luckert, Mohr, Daba, Warkentin, Hadinezhad, Barlow, Hou, Lefebvre, Golshani, Cober and Samanfar</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Soybean is an important global source of plant-based protein. A persistent trend has been observed over the past two decades that soybeans grown in western Canada have lower seed protein content than soybeans grown in eastern Canada. In this study, 10 soybean genotypes ranging in average seed protein content were grown in an eastern location (control) and three western locations (experimental) in Canada. Seed protein and oil contents were measured for all lines in each location. RNA-sequencing and differential gene expression analysis were used to identify differentially expressed genes that may account for relatively low protein content in western-grown soybeans. Differentially expressed genes were enriched for ontologies and pathways that included amino acid biosynthesis, circadian rhythm, starch metabolism, and lipid biosynthesis. Gene ontology, pathway mapping, and quantitative trait locus (QTL) mapping collectively provide a close inspection of mechanisms influencing nitrogen assimilation and amino acid biosynthesis between soybeans grown in the East and West. It was found that western-grown soybeans had persistent upregulation of asparaginase (an asparagine hydrolase) and persistent downregulation of asparagine synthetase across 30 individual differential expression datasets. This specific difference in asparagine metabolism between growing environments is almost certainly related to the observed differences in seed protein content because of the positive correlation between seed protein content at maturity and free asparagine in the developing seed. These results provided pointed information on seed protein-related genes influenced by environment. This information is valuable for breeding programs and genetic engineering of geographically optimized soybeans.</p>
</abstract>
<kwd-group>
<kwd>RNA-seq</kwd>
<kwd>differential expression</kwd>
<kwd>soybean</kwd>
<kwd>asparagine</kwd>
<kwd>seed protein</kwd>
<kwd>amino acid metabolism</kwd>
</kwd-group>
<counts>
<fig-count count="7"/>
<table-count count="1"/>
<equation-count count="0"/>
<ref-count count="67"/>
<page-count count="20"/>
<word-count count="12068"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-in-acceptance</meta-name>
<meta-value>Functional and Applied Plant Genomics</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1" sec-type="intro">
<label>1</label>
<title>Introduction</title>
<p>Soybean (<italic>Glycine max</italic> [L.] Merr.) is one of the most important legume crops worldwide for use as human food and livestock feed. Soybean seeds contain the highest protein of any legume, which makes protein content a key quality attribute (<xref ref-type="bibr" rid="B43">Natarajan et&#xa0;al., 2013</xref>; <xref ref-type="bibr" rid="B32">Huang et&#xa0;al., 2019</xref>). In symbiosis with rhizobia, soybean fixes atmospheric nitrogen into more biologically available forms of nitrogen. Nitrogen fixation gives soybeans a valuable role in sustainable agricultural practices by reducing the need for nitrogen fertilizers and reducing incomplete nitrogen conversion/capture, which pollutes the surrounding environment (air, soil, and water). As the global population rises, strategic agricultural planning requires optimization of crops for different environmental conditions in order to produce adequate yields with acceptable levels of high-quality seed protein.</p>
<p>For more than two decades, observations have been made that soybeans grown in western Canada have lower (~1%&#x2013;5%) seed protein content than eastern-grown soybeans. In 2022, the average eastern soybean protein content was 40.3%, while the western-grown soybeans had an average protein content of 38.9% (<xref ref-type="bibr" rid="B10">Canadian Grain Commission, 2022</xref>). In general, western Canadian soybean-growing regions have lower seasonal precipitation, sandier soils, shorter growing seasons with a longer photoperiod, and cooler temperatures, all of which contribute to reduced seed quality and are attributed to difficulties in successfully growing soybeans. Seed protein and oil contents are complex quantitatively inherited traits that are influenced by the combination of genotype and environment (<xref ref-type="bibr" rid="B63">Wang et&#xa0;al., 2019</xref>). Soybeans from western Canada have been observed to spend more time in the vegetative stages of development compared to eastern-grown counterparts and less time dedicated to flowering and seed development (<xref ref-type="bibr" rid="B44">Ort et&#xa0;al., 2022</xref>). Seed protein content is a key factor in soybean seed quality measures; profits are significantly impacted for farmers who grow soybeans in suboptimal environmental conditions. Further, as climates change and populations increase, it is of great economic and agricultural importance to make better use of the western and northern growing regions of Canada.</p>
<p>Soybean agronomic productivity is measured in part by seed composition, specifically regarding the total seed content of two major seed storage biomolecules: protein and oil. Generally, protein and oil contents have an inverse relationship in soybean seed; as oil increases, protein decreases, and vice versa (<xref ref-type="bibr" rid="B8">Breene et&#xa0;al., 1988</xref>; <xref ref-type="bibr" rid="B13">Clemente and Cahoon, 2009</xref>). Expression of genes involved in seed protein and oil is most highly expressed during the middle and late stages of development (<xref ref-type="bibr" rid="B53">Severin et&#xa0;al., 2010</xref>). Genes involved in fatty acid synthesis and elongation (<italic>fad2</italic>, <italic>lox</italic>, and <italic>kcs</italic>) contribute to the differences seen between high-protein-low-oil and low-protein-high-oil cultivars, while seed protein content differences are influenced by transcription factors (including <italic>abi3</italic> and <italic>lec2</italic>), and sugar transporter <italic>SWEET10a</italic> plays a role in both protein and oil accumulation (<xref ref-type="bibr" rid="B61">Wang et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B47">Peng et&#xa0;al., 2021</xref>). In a recent study, lipid and carbohydrate metabolism were found to be differentially expressed high-protein and high 11S soybeans in comparison to their low-protein and low-11S counterparts grown in the same conditions (<xref ref-type="bibr" rid="B30">Hooker et al., 2023</xref>). High protein requires the plant to direct a significant amount of nitrogen assimilates to the seeds. Total amino acid content is positively correlated with seed protein content in soybeans (<xref ref-type="bibr" rid="B67">Zhang et&#xa0;al., 2018</xref>). It has been observed that protein content at the time of soybean seed maturity is positively correlated with free asparagine in developing seeds (<xref ref-type="bibr" rid="B29">Hern&#xe1;ndez-Sebasti&#xe0; et&#xa0;al., 2005</xref>; <xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>). This relationship has also been observed in other crops, such as barley and maize (<xref ref-type="bibr" rid="B18">Dembinski and Bany, 1991</xref>; <xref ref-type="bibr" rid="B39">Lohaus et&#xa0;al., 1998</xref>). Soybean seed protein accumulation is controlled in part by the biosynthesis of nitrogenous assimilates in source leaves. Biologically available inorganic nitrogen in the soil (ammonium NH<sub>4</sub><sup>+</sup>, nitrite NO<sub>2</sub><sup>&#x2212;</sup>, and nitrate NO<sub>3</sub><sup>&#x2212;</sup>) must be reduced to ammonia (NH<sub>3</sub>) before assimilation into amino acids (<xref ref-type="bibr" rid="B35">Lea and Miflin, 1980</xref>; <xref ref-type="bibr" rid="B36">Lea et&#xa0;al., 1990</xref>). However, there is a gap in understanding the mechanisms underlying the accumulation of nitrogen assimilates in the developing seed; it is unclear whether nitrogen assimilate supply is directed by the mother plant or if the developing seed has an intrinsic capacity for storage protein synthesis (<xref ref-type="bibr" rid="B29">Hern&#xe1;ndez-Sebasti&#xe0; et&#xa0;al., 2005</xref>). In large-seeded species like beans (<italic>Phaseolus limensis</italic> L.), seed size is sufficient enough that they have a vascular bundle, which allows for the direct distribution of nutrients from the mother plant to the developing seed (<xref ref-type="bibr" rid="B56">Vinogradova and Falaleev, 2012</xref>); however, further exploration into these processes in soybeans is needed.</p>
<p>RNA-sequencing (RNA-seq) and differential expression (DE) are powerful tools for functional genomics and transcriptomics. Comparing gene expression between two genetically identical samples in two different environmental conditions allows for a snapshot of the active and inactive genes directly influenced by the environment. Downstream analyses of the resulting DE genes give valuable information on the pathways and systems that are influenced by a given environment. Gene ontology (GO) and pathway mapping databases are regularly updated with novel and/or more robust information that can process large lists of genes to provide multi-perspective functional analysis. Quantitative trait locus (QTL) analysis uses variable quantitative traits and genotypic information to make correlations between the two. QTL mapping is useful for identifying molecular loci influential of a given biological pathway and offers information about locus location and linkage. Collectively, there are over 550 protein and oil QTLs known and available in SoyBase (<ext-link ext-link-type="uri" xlink:href="http://www.soybase.org">www.soybase.org</ext-link>) distributed over all 20 chromosomes, but there are higher proportions falling on chromosomes 5, 15, and 20 (<xref ref-type="bibr" rid="B63">Wang et&#xa0;al., 2019</xref>). By uncovering DE genes with key functional roles in seed protein and/or oil development, avenues for genetic engineering of soybeans become more effective for breeding agriculturally sustainable soybean crops. It is hypothesized that soybeans grown in western growing regions are differentially expressing some of their seed protein-related genes when compared to those grown in the East. The objective of this study was to investigate the differential gene expression between soybeans grown in East and West Canada to uncover the key metabolic pathways potentially influencing seed protein content. To do this, RNA-seq and DE data were collected in 2019 spanning 10 soybean varieties, three experimental locations (West), and one control location (East).</p>
</sec>
<sec id="s2">
<label>2</label>
<title>Methods</title>
<sec id="s2_1">
<label>2.1</label>
<title>Soybean lines</title>
<p>Ten soybean genotypes were selected as a representation of the range of seed protein content observed in Canada. Soybean lines are listed from the lowest to highest seed protein content, with line 1 having the lowest average protein content and line 10 having the highest average protein content. The soybean lines used in this study were all developed at the Ottawa Research and Development Centre by Agriculture and Agri-Food Canada: X5583-1-041-5-5 (line 1), AC Harmony (<xref ref-type="bibr" rid="B57">Voldeng et&#xa0;al., 1996a</xref>) (line 2), AAC Halli (line 3), 90A01 (<xref ref-type="bibr" rid="B14">Cober et&#xa0;al., 2006</xref>) (line 4), Maple Amber (line 5), OT13-08 (line 6), OT14-03 (line 7), AAC Springfield (line 8), Jari (line 9), and AC Proteus (<xref ref-type="bibr" rid="B58">Voldeng et&#xa0;al., 1996b</xref>) (line 10).</p>
</sec>
<sec id="s2_2">
<label>2.2</label>
<title>Planting and growth</title>
<p>Planting was performed in 2019 in replicated trials across four locations: Ottawa Ontario (latitude 45.39&#xb0;, longitude &#x2212;75.72&#xb0;), Morden Manitoba (49.18&#xb0;, &#x2212;98.08&#xb0;), Brandon Manitoba (49.86&#xb0;, &#x2212;99.98&#xb0;), and Saskatoon Saskatchewan (52.15&#xb0;, &#x2212;106.57&#xb0;). Seeds were planted in quadruplicate at the mid-end of May in a 4 &#xd7; 5 rectangular lattice arrangement at a density of 50 seeds per m<sup>2</sup>, and appropriate crop management practices were taken at each site. For additional information on planting, see <xref ref-type="bibr" rid="B15">Cober et&#xa0;al. (2023)</xref>.</p>
</sec>
<sec id="s2_3">
<label>2.3</label>
<title>Tissue collection and seed composition assessment</title>
<p>Young trifoliate leaf tissue was collected in triplicate from soybeans at the R5 (<xref ref-type="bibr" rid="B46">Pedersen and Licht, 2014</xref>) stage of maturity from otherwise healthy-looking plants. Tissue was flash-frozen in liquid nitrogen in the field immediately upon harvest, and samples were stored at &#x2212;80&#xb0;C. Western samples were shipped overnight on dry ice and immediately stored at &#x2212;80&#xb0;C until RNA extraction. From each plot, measurements for total seed protein and oil contents were performed using an Infratec 1241 Grain Analyzer (FOSS North America, Eden Prairie, MN, USA) at the Agriculture and Agri-Food Canada Ottawa Research and Development Centre. For additional phenotypic information, see [25].</p>
</sec>
<sec id="s2_4">
<label>2.4</label>
<title>RNA extraction</title>
<p>RNA extractions using SPLIT Total mRNA Extraction Kit (Lexogen, Vienna, Austria) were performed on approximately 200 mg of crushed leaf tissue from each sample according to the manufacturer&#x2019;s instructions. RNA quality was tested using a NanoDrop&#x2122; 2000 Spectrophotometer (Thermo Fisher Scientific, Waltham, MA, USA), agarose gel electrophoresis (1%), TapeStation 4200 RNA ScreenTape (Agilent, Santa Clara, CA, USA), and 2100 Bioanalyzer (Agilent, Santa Clara, CA, USA) at G&#xe9;nome Qu&#xe9;bec (Montr&#xe9;al, QC, Canada) and the Ottawa Research and Development Centre (Ottawa, ON, Canada). RNA integrity number (RIN) values of at least 6.5 and a Q30 score of at least 36 were selected for library preparation. Spike-in RNA variants (SIRVs) (Lexogen, Vienna, Austria) were integrated into the RNA samples as controls to monitor and compare key parameters (such as sensitivity and quantification); the E0 SIRV mix was used, containing 69 different isoform variants with known sequences at equal molar concentrations.</p>
</sec>
<sec id="s2_5">
<label>2.5</label>
<title>RNA-seq library preparation, alignment, read mapping, read counting, and DE</title>
<p>Paired-end sequencing was carried out using the Illumina HiSeq 4000 platform (Illumina, San Diego, CA, USA) at G&#xe9;nome Qu&#xe9;bec to create cDNA libraries for each sample. RNA-seq data were assessed using dupRadar (<xref ref-type="bibr" rid="B51">Sayols et&#xa0;al., 2016</xref>) (v3.16, Biberach an der Ri&#xdf;, Germany; Bioconductor, R) for duplication rate quality control. Normalization of reads was carried out at the individual sample level using edgeR (<xref ref-type="bibr" rid="B50">Robinson et&#xa0;al., 2010</xref>) (v3.16, Parkville, Victoria, Australia). Exploratory data analysis of normalized reads was performed using R.</p>
<p>QualiMap (<xref ref-type="bibr" rid="B22">Garc&#xed;a-Alcalde et&#xa0;al., 2012</xref>) (v2.2.1, Berlin, Germany) was used as a quality control step for sequence data feature alignment (genes and transcripts). Preseq (<xref ref-type="bibr" rid="B16">Daley et&#xa0;al., 2020</xref>) (v3.1.1, Los Angeles, CA, USA; Bioconductor, R) was used to estimate the number of distinct reads from each RNA-seq library. RSeQC (<xref ref-type="bibr" rid="B62">Wang et&#xa0;al., 2012</xref>) (v4.0.0, Nanjing, China; Bioconductor, R) was the primary tool used for comprehensive evaluation of the RNA-seq read data through the calculation of semantic read distribution of a sample, distance between reads, duplication presence, and junction saturation.</p>
<p>The Canadian Centre for Computational Genomics uses an in-house framework program, GenPipes (<xref ref-type="bibr" rid="B7">Bourgey et&#xa0;al., 2019</xref>), to perform the following major processing steps. Trimmomatic (<xref ref-type="bibr" rid="B6">Bolger et&#xa0;al., 2014</xref>) (v0.36, J&#xfc;lich, Germany) was used to remove adaptor sequences and low quality score bases (phred score &lt;30). Trimmed reads were aligned to the soybean genome (Glycine_max_v2.1, INSDC Assembly GCA_000004515.4, Jul 2018), using STAR (<xref ref-type="bibr" rid="B19">Dobin et&#xa0;al., 2013</xref>) (v2.7.7a, Menlo Park, CA, USA) under the command &#x2013;runMode alignReads after generating index files from the aforementioned genome. HTSeq (<xref ref-type="bibr" rid="B3">Anders et&#xa0;al., 2015</xref>) (v0.12.3, Heidelberg, Germany) was used to obtain read counts using the following options: &#x201c;-m intersection-nonempty&#x201d;.</p>
</sec>
<sec id="s2_6">
<label>2.6</label>
<title>Differential gene expression analysis and candidate gene identification</title>
<p>DE analysis was performed using DESeq2 (<xref ref-type="bibr" rid="B40">Love et&#xa0;al., 2014</xref>) (v3.16, Heidelberg, Germany) with negative Binomial GLM fitting and Wald statistics: nbinomWaldTest. To transform expression data to be expressed as log<sub>2</sub>FC, &#x201c;ashr&#x201d; (<xref ref-type="bibr" rid="B54">Stephens, 2017</xref>) was used. All datasets were trimmed to an adjusted p-value &lt;0.01. For DE analysis, Ottawa samples were used as the control, and the three western locations were each used as the experimental data; the log<sub>2</sub> fold change (log<sub>2</sub>FC) difference in expression data reflects a change occurring in the western-grown relative to eastern-grown cultivars. Identical genotypes were compared for each DE analysis, and comparisons were not made across different genotypes.</p>
<p>All DE datasets were amended with the corresponding information from the SoyBase Genome Annotation Source v2.0 (<ext-link ext-link-type="uri" xlink:href="https://soybase.org/genomeannotation/">https://soybase.org/genomeannotation/</ext-link>), which includes annotation data from BLASTP, TAIR10, GO, Panther, PFAM, and KOG for all genes. In this study, both top-down and bottom-up analyses were used to assess DE genes between East and West. &#x201c;Top-down&#x201d; and &#x201c;bottom-up&#x201d; are descriptive terms for the direction of data analysis. To clarify, the top-down analysis uses the holistic dataset and investigates the DE genes from a bird&#x2019;s-eye perspective without any specific functional selection&#x2014;i.e., which genes (and their ontologies) are DE between East and West at the given cut-off criteria (p-value &lt;0.01, |log2FC| &#x2265; 1.5). The top-down approach was used to holistically search the DE data for the most consistent DE genes between East and West; DE was cross-compared across all 30 datasets for most of the consistently (30/30 datasets) DE genes. The bottom-up analysis describes a different approach to the data: an investigation in which we search only for genes involved in one specific pathway of interest (the Asp-Ala-Glu pathway). The bottom-up approach was used to search through the DE data to identify genes with &#x201c;asparagine&#x201d;, &#x201c;aspartate&#x201d;, &#x201c;alanine&#x201d;, &#x201c;glutamate&#x201d;, and &#x201c;oxaloacetate&#x201d; as a component of their annotation (BLASTP, TAIR10, GO, Panther, PFAM, or KOG). With the use of a short bash script, all 30 DE datasets were searched for any gene with these keyword identifiers and were short-listed for pathway-specific analysis. The purpose of this bottom-up analysis was to provide a comprehensive investigation of DE genes within this pathway and provide insight into the underlying molecular mechanisms influencing low seed protein phenotypes in western Canada. A p-value &lt;0.01 and a log<sub>2</sub>FC change in expression of at least 1.5 (|log2FC| &#x2265; 1.5) were considered significantly DE for both the top-down and bottom-up approaches. Genes that are DE in 15 (50%) or more datasets were selected for downstream analysis. Because the SoyBase annotation database uses BLAST descriptions from 2014, an up-to-date BLAST search was carried out on all resultant top-down and bottom-up genes and is included alongside the SoyBase annotation output in <xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref> and <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>.</p>
</sec>
<sec id="s2_7">
<label>2.7</label>
<title>GO analysis</title>
<p>For the top-down analyses, GO term enrichment was assessed using the SoyBase GO Term Enrichment Tool (<ext-link ext-link-type="uri" xlink:href="https://soybase.org/goslimgraphic_v2/dashboard.php">https://soybase.org/goslimgraphic_v2/dashboard.php</ext-link>) for the genes commonly upregulated and commonly downregulated across all 30 DE datasets. Enrichment was calculated from the ratio of expressed DE genes for a particular GO term to the expected number of DE genes for said term based on the total number of known associated genes in the full GO database.</p>
</sec>
<sec id="s2_8">
<label>2.8</label>
<title>Heatmapping</title>
<p>Heatmapping was performed using Heatmapper (<xref ref-type="bibr" rid="B5">Babicki et&#xa0;al., 2016</xref>) using log-normalized read counts across all samples in this study calculated using R. Clustering was calculated using average linkage, and distance matrices were calculated using Pearson&#x2019;s coefficient. The row Z-score normalizes expression data to improve visualization of heatmap data trends; this score is calculated by (gene expression value in sample of interest) &#x2212; (mean expression across all samples)/(standard deviation) (<xref ref-type="bibr" rid="B2">Anders and Huber, 2010</xref>).</p>
</sec>
<sec id="s2_9">
<label>2.9</label>
<title>KEGG analysis</title>
<p>Gene IDs were converted to their corresponding National Center for Biotechnology Information (NCBI) ID numbers, which were then mapped using the Kyoto Encyclopedia of Genes and Genomes (KEGG) release v106.0 (<ext-link ext-link-type="uri" xlink:href="https://www.kegg.jp/">https://www.kegg.jp/</ext-link>) for pathway enrichment using the Mapper Search functions, with <italic>Glycine max</italic> (gmx) as the organism identifier.</p>
</sec>
<sec id="s2_10">
<label>2.10</label>
<title>QTL analysis</title>
<p>Chromosome positioning data for all 20 <italic>G. max</italic> chromosomes were extracted from SoyBase GWAS-based QTL database for all seed protein and oil QTLs (<ext-link ext-link-type="uri" xlink:href="https://soybase.org/GWAS/list.php">https://soybase.org/GWAS/list.php</ext-link>). The positional information for each of the 59 bottom-up genes of interest was also extracted from SoyBase and used to determine which genes fall within large QTL regions or in very close proximity to single-point QTLs. With the use of MapChart v2.32 (<xref ref-type="bibr" rid="B59">Voorrips, 2002</xref>), QTL and gene positional information were mapped to each chromosome and color-coded to be most informative.</p>
</sec>
</sec>
<sec id="s3" sec-type="results">
<label>3</label>
<title>Results</title>
<sec id="s3_1">
<label>3.1</label>
<title>RNA-seq analyses</title>
<p>There was a total of 4,047,045,039 reads over all 87 RNA-seq datasets (10 lines, four locations, and three replicates per line). Missing replicates are Saskatoon line 1 replicate 1, Saskatoon line 4 replicate 2, and Saskatoon line 10 replicate 3, which did not pass quality control (QC) upon repeated attempts. <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref> summarizes the average seed protein and oil contents as a percentage of the total seed content. Included in <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref> is the cumulative read depth across the three replicates per sample. <xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref> shows the principal component analysis (PCA) of the transcriptome data for each replicate in East and West locations, organized by line (1&#x2013;10). Across all samples, PC1 described a 66% variance, and PC2 described a 7% variance (<xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref>). From <xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref>, there is a clear distinction between East (cyan) and West (green, purple, and red) RNA-seq variability, indicating the RNA-seq data from the East are highly different than the data from the West, and the data from the three West locations cluster closely together, suggesting similar variability.</p>
<table-wrap id="T1" position="float">
<label>Table&#xa0;1</label>
<caption>
<p>Average seed protein and oil contents (given in percentage of total seed content at 13% moisture) and RNA-seq read depth for the three replicates for each soybean genotype in East and West locations in Canada in 2019.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" rowspan="3" align="center">Line</th>
<th valign="top" colspan="3" align="center">East</th>
<th valign="top" colspan="9" align="center">West</th>
</tr>
<tr>
<th valign="top" colspan="3" align="center">Ottawa</th>
<th valign="top" colspan="3" align="center">Morden</th>
<th valign="top" colspan="3" align="center">Brandon</th>
<th valign="top" colspan="3" align="center">Saskatoon</th>
</tr>
<tr>
<th valign="top" align="left">Protein</th>
<th valign="top" align="left">Oil</th>
<th valign="top" align="left">Read depth</th>
<th valign="top" align="left">Protein</th>
<th valign="top" align="left">Oil</th>
<th valign="top" align="left">Read depth</th>
<th valign="top" align="left">Protein</th>
<th valign="top" align="left">Oil</th>
<th valign="top" align="left">Read depth</th>
<th valign="top" align="left">Protein</th>
<th valign="top" align="left">Oil</th>
<th valign="top" align="left">Read depth</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">
<bold>
<italic>1</italic>
</bold>
</td>
<td valign="top" align="left">38.9</td>
<td valign="top" align="left">23.0</td>
<td valign="top" align="left">116,820,021</td>
<td valign="top" align="left">37.3</td>
<td valign="top" align="left">22.2</td>
<td valign="top" align="left">92,391,979</td>
<td valign="top" align="left">36.3</td>
<td valign="top" align="left">21.2</td>
<td valign="top" align="left">95,080,016</td>
<td valign="top" align="left">37.7</td>
<td valign="top" align="left">19.3</td>
<td valign="top" align="left">69,995,998</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>2</italic>
</bold>
</td>
<td valign="top" align="left">37.5</td>
<td valign="top" align="left">23.4</td>
<td valign="top" align="left">92,142,359</td>
<td valign="top" align="left">35.6</td>
<td valign="top" align="left">22.8</td>
<td valign="top" align="left">84,842,378</td>
<td valign="top" align="left">35.2</td>
<td valign="top" align="left">21.3</td>
<td valign="top" align="left">119,248,964</td>
<td valign="top" align="left">37.5</td>
<td valign="top" align="left">19.3</td>
<td valign="top" align="left">104,853,205</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>3</italic>
</bold>
</td>
<td valign="top" align="left">38.9</td>
<td valign="top" align="left">22.1</td>
<td valign="top" align="left">93,160,198</td>
<td valign="top" align="left">38.2</td>
<td valign="top" align="left">21.7</td>
<td valign="top" align="left">100,375,142</td>
<td valign="top" align="left">39.0</td>
<td valign="top" align="left">19.7</td>
<td valign="top" align="left">89,082,785</td>
<td valign="top" align="left">38.9</td>
<td valign="top" align="left">18.9</td>
<td valign="top" align="left">118,090,282</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>4</italic>
</bold>
</td>
<td valign="top" align="left">40.9</td>
<td valign="top" align="left">21.5</td>
<td valign="top" align="left">92,490,279</td>
<td valign="top" align="left">39.5</td>
<td valign="top" align="left">21.3</td>
<td valign="top" align="left">99,640,813</td>
<td valign="top" align="left">40.2</td>
<td valign="top" align="left">20.1</td>
<td valign="top" align="left">97,743,936</td>
<td valign="top" align="left">39.8</td>
<td valign="top" align="left">18.5</td>
<td valign="top" align="left">92,095,914</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>5</italic>
</bold>
</td>
<td valign="top" align="left">40.9</td>
<td valign="top" align="left">22.1</td>
<td valign="top" align="left">119,385,027</td>
<td valign="top" align="left">40.1</td>
<td valign="top" align="left">21.6</td>
<td valign="top" align="left">103,073,578</td>
<td valign="top" align="left">39.3</td>
<td valign="top" align="left">20.3</td>
<td valign="top" align="left">92,696,520</td>
<td valign="top" align="left">40.1</td>
<td valign="top" align="left">18.7</td>
<td valign="top" align="left">89,401,807</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>6</italic>
</bold>
</td>
<td valign="top" align="left">42.3</td>
<td valign="top" align="left">21.9</td>
<td valign="top" align="left">123,431,881</td>
<td valign="top" align="left">41.6</td>
<td valign="top" align="left">21.7</td>
<td valign="top" align="left">93,502,218</td>
<td valign="top" align="left">41.1</td>
<td valign="top" align="left">20.8</td>
<td valign="top" align="left">90,866,478</td>
<td valign="top" align="left">40.7</td>
<td valign="top" align="left">20.0</td>
<td valign="top" align="left">136,867,397</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>7</italic>
</bold>
</td>
<td valign="top" align="left">41.7</td>
<td valign="top" align="left">20.4</td>
<td valign="top" align="left">104,008,337</td>
<td valign="top" align="left">42.1</td>
<td valign="top" align="left">20.0</td>
<td valign="top" align="left">112,032,906</td>
<td valign="top" align="left">41.9</td>
<td valign="top" align="left">18.0</td>
<td valign="top" align="left">106,855,434</td>
<td valign="top" align="left">40.8</td>
<td valign="top" align="left">18.1</td>
<td valign="top" align="left">136,026,735</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>8</italic>
</bold>
</td>
<td valign="top" align="left">45.1</td>
<td valign="top" align="left">18.5</td>
<td valign="top" align="left">97,175,295</td>
<td valign="top" align="left">43.2</td>
<td valign="top" align="left">19.5</td>
<td valign="top" align="left">98,480,553</td>
<td valign="top" align="left">42.9</td>
<td valign="top" align="left">17.9</td>
<td valign="top" align="left">101,918,836</td>
<td valign="top" align="left">45.2</td>
<td valign="top" align="left">16.6</td>
<td valign="top" align="left">93,969,604</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>9</italic>
</bold>
</td>
<td valign="top" align="left">44.3</td>
<td valign="top" align="left">19.0</td>
<td valign="top" align="left">101,408,312</td>
<td valign="top" align="left">42.3</td>
<td valign="top" align="left">19.7</td>
<td valign="top" align="left">96,174,311</td>
<td valign="top" align="left">41.3</td>
<td valign="top" align="left">18.3</td>
<td valign="top" align="left">114,640,094</td>
<td valign="top" align="left">43.5</td>
<td valign="top" align="left">17.3</td>
<td valign="top" align="left">103,655,546</td>
</tr>
<tr>
<td valign="top" align="center">
<bold>
<italic>10</italic>
</bold>
</td>
<td valign="top" align="left">47.5</td>
<td valign="top" align="left">16.7</td>
<td valign="top" align="left">101,581,019</td>
<td valign="top" align="left">46.9</td>
<td valign="top" align="left">16.7</td>
<td valign="top" align="left">115,579,425</td>
<td valign="top" align="left">46.2</td>
<td valign="top" align="left">15.9</td>
<td valign="top" align="left">103,436,842</td>
<td valign="top" align="left">47.0</td>
<td valign="top" align="left">16.3</td>
<td valign="top" align="left">52,822,615</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="f1" position="float">
<label>Figure&#xa0;1</label>
<caption>
<p>PCA for each soybean genotype (number at top of each panel) at each location based on total RNA-seq variance after removal of outliers and normalization. Soybean genotype number is represented above each corresponding PCA plot. PC1 is on the x-axis, and PC2 is on the y-axis. Gray points represent all the data points in other lines. Replicates within each line (R1, R2, and R3) are labeled at their corresponding points. PCA, principal component analysis (<xref ref-type="bibr" rid="B31">Hooker et al., 2022</xref>).</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-14-1260393-g001.tif"/>
</fig>
</sec>
<sec id="s3_2">
<label>3.2</label>
<title>Top-down approach to DE analysis</title>
<sec id="s3_2_1">
<label>3.2.1</label>
<title>Upregulated genes</title>
<p>For the top-down analysis, genes identified to be significantly DE (|log<sub>2</sub>FC| &#x2265; 1.5, p-value &lt;0.01) across all 30 datasets with the same orientation (up- or downregulated in the West compared to the East) were considered. For all DE genes for each line in each location, including unique IDs and commonly DE genes, see <xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>. The top-down analysis from Brandon had a total of 34,984 instances of upregulated genes across all 10 genotypes, composed of 8,652 unique gene IDs, of which 774 were commonly upregulated across all 10 genotypes (datasets) (<xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). The Morden top-down analysis found a total of 38,866 instances of upregulated genes across all 10 lines, composed of a total of 8,620 unique gene IDs, 1,226 of which are commonly upregulated across all 10 lines in Morden (<xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). The Saskatoon top-down analysis had a total of 52,100 upregulated genes across the 10 datasets, composed of a total of 10,812 unique IDs, of which 1,679 were commonly upregulated across all 10 lines in Saskatoon (<xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). In total, 514 genes were commonly upregulated across all 30 East <italic>vs.</italic> West DE datasets (10 lines, three West locations) (<xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). <xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2A</bold>
</xref> shows a Venn diagram of the genes commonly upregulated across all lines; <xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2B</bold>
</xref> shows the commonly downregulated genes. The values given in the exterior &#x201c;petals&#x201d; of the Venn diagram represent the number of genes that were found to be either upregulated (<xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2A</bold>
</xref>) or downregulated (<xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2B</bold>
</xref>) in the individual East <italic>vs.</italic> West DE analyses, which were used to find common DE genes; this was performed to circumvent the infeasibility of presenting a 30-way Venn diagram with all possible combinations of line-location overlapping genes.</p>
<fig id="f2" position="float">
<label>Figure&#xa0;2</label>
<caption>
<p>Venn diagrams of DE soybean genes in western locations relative to east. Modified Venn diagram of the number of genes in top-down DE analyses for each western location. Small exterior petals represent the number of <bold>(A)</bold> upregulated or <bold>(B)</bold> downregulated genes in each line-location pairwise comparison (p-value &lt;0.01, log<sub>2</sub>FC 1.5). &#x201c;M&#x201d; (blue) represents Morden, &#x201c;B&#x201d; (red) represents Brandon, and &#x201c;S&#x201d; (green) represents Saskatoon. Gene IDs commonly up- or downregulated across all lines 1&#x2013;10 per location were used to construct the Venn diagram. Common DE genes were found using the VLOOKUP function in MS Excel, and the Venn diagram was created using <uri xlink:href="https://bioinformatics.psb.ugent.be/webtools/Venn/">https://bioinformatics.psb.ugent.be/webtools/Venn/</uri>. DE, differential expression.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-14-1260393-g002.tif"/>
</fig>
</sec>
<sec id="s3_2_2">
<label>3.2.2</label>
<title>Downregulated genes</title>
<p>In total, there were 29,826 instances of downregulation across all 10 lines in Morden, made up of 6,475 unique gene IDs and 956 genes commonly downregulated across the 10 lines (<xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). In Brandon, there were 25,505 instances of downregulation across all 10 lines, composed of 6,469 unique gene IDs, of which 619 were commonly downregulated across all 10 lines (<xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). There were 34,978 instances of downregulation in Saskatoon across all 10 DE datasets, composed of 6,931 unique gene IDs, and 1,328 of those genes were commonly downregulated across all 10 lines. There were 415 genes commonly downregulated across all 30 East <italic>vs.</italic> West datasets (<xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2B</bold>
</xref>; <xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>).</p>
</sec>
<sec id="s3_2_3">
<label>3.2.3</label>
<title>Gene ontology</title>
<p>After GO enrichment using the SoyBase GO Term Enrichment Tool, there were 730 GO terms (biological process (BP) and molecular function (MF)) associated with the genes consistently downregulated in the West and 815 GO terms associated with the genes consistently upregulated in the West (<xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). <xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref> shows the most highly enriched GO terms (BP and MF) across the genes consistently DE in the West across all lines. Bubble size indicates the number of DE genes in our list that are associated with a particular term. Enrichment was calculated using the proportion of the number of DE genes observed to be associated with a term divided by the number of genes expected to be among a list of the query size. GO terms graphed in <xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref> were selected based on an enrichment score over 2 (i.e., enriched by 100% or twofold) and at least five genes in our DE data with a given term included in their annotations. Terms are listed in order from the highest number of DE genes per term to the lowest (five genes, minimum). <xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Table&#xa0;1</bold>
</xref> summarizes the GO enrichment data for the persistently up- and downregulated genes. Included among the most enriched genes upregulated in the West are cytokinesis by cell plate formation (GO:0000911), spindle assembly (GO:0051225), microtubule motor activity (GO:0003777), response to UV (GO:0009411), long day photoperiodism (flowering) (GO:0048574), and cyclin-dependent protein serine/threonine kinase regulator activity (GO:0016538) (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3A</bold>
</xref>). Other noteworthy GOs from the upregulated genes include lipid-related ontologies, fatty acid biosynthetic process (GO:0006633), lipid transport (GO:0006869), and lipid binding (GO:0008289). Circadian rhythm (GO:0007623) was among the most enriched genes consistently up- and downregulated in the West (<xref ref-type="fig" rid="f3">
<bold>Figures&#xa0;3A, B</bold>
</xref>). Regulation of circadian rhythm (GO:0042752), secondary metabolic process (GO:0019748), starch metabolic process (GO:0005982), and starch biosynthetic process (GO:0019252) were among the topmost enriched GO terms from the list of genes downregulated in the West. Cellular amino acid biosynthetic process (GO:2000282), asparagine biosynthesis, aromatic amino acid family metabolic process (GO:0009072), and maltose metabolic process (GO:0000023) were noteworthy terms among the highly enriched downregulated genes (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3B</bold>
</xref>).</p>
<fig id="f3" position="float">
<label>Figure&#xa0;3</label>
<caption>
<p>Enriched <bold>(A)</bold> upregulated and <bold>(B)</bold> downregulated BP and MF GO terms for top-down analysis of DE soybean genes between East and West. Enrichment was calculated by taking the proportion of the number of DE genes associated with a term and the expected number of genes associated with a term. Represented in the graphs are the GO terms with enrichment values of at least 2 (overrepresented by 100% or twofold) with a minimum of five expressed genes with GO terms in list. BP, biological process; MF, molecular function; GO, gene ontology; DE, differential expression.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-14-1260393-g003.tif"/>
</fig>
</sec>
<sec id="s3_2_4">
<label>3.2.4</label>
<title>KEGG pathway enrichment</title>
<p>Gene IDs from up- and downregulated top-down DE analyses were converted to their NCBI gene ID number and then mapped to the <italic>G. max</italic> database (gmx) within the KEGG for pathway enrichment. Of the topmost enriched pathways, 109 genes mapped to the broad <italic>G. max</italic> metabolic pathway (gmx01100), 68 genes mapped to the biosynthesis of secondary metabolites pathway (gmx01110), 22 genes mapped to motor proteins (gmx04814), 13 genes mapped to plant hormone signal transduction (gmx04075), 12 genes mapped to circadian rhythm &#x2013; plant (gmx04712), 11 genes mapped to carbon metabolism (gmx01200), and 10 genes mapped to the biosynthesis of cofactors (gmx01240). Other pathways of note include aromatic amino acid (phenylalanine, tyrosine, and tryptophan) biosynthesis (gmx00400; three genes), fatty acid metabolism (gmx01212; three genes), and sulfur-containing amino acid (cysteine and methionine) metabolism (gmx00270; three genes). The full list of enriched pathways and the genes that map to each are in <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>. <xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4</bold>
</xref> shows the enriched pathways within the broad <italic>G. max</italic> metabolic pathway (gmx01100); the red highlight indicates pathways DE between East and West grown soybeans.</p>
<fig id="f4" position="float">
<label>Figure&#xa0;4</label>
<caption>
<p>KEGG pathway enrichment of the up- and downregulated genes across all 30 DE datasets (|log<sub>2</sub>FC| &#x2265; 1.5, p-value &lt;0.01<italic>) across all known soybean metabolic pathways (gmx01100)</italic>. Green highlight indicates known pathways in soybeans, and red highlight indicates pathways DE between eastern- and western-grown soybeans. For a larger image of this map, see <xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Figure&#xa0;1</bold>
</xref>. KEGG, Kyoto Encyclopedia of Genes and Genomes; DE, differential expression.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-14-1260393-g004.tif"/>
</fig>
</sec>
</sec>
<sec id="s3_3">
<label>3.3</label>
<title>Bottom-up approach to DE analysis</title>
<p>The top-down analysis showed significant enrichment of biosynthesis of secondary metabolites (gmx01110), which led to the downstream investigation of select sub-pathways. Among these sub-pathways, one pathway of interest was the alanine, aspartate, and glutamate (Ala-Asp-Glu) metabolism pathway (gmx00250), which was of particular interest because of the known relationship between nitrate assimilation during development and seed protein content at maturity (<xref ref-type="bibr" rid="B29">Hern&#xe1;ndez-Sebasti&#xe0; et&#xa0;al., 2005</xref>; <xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>). Asparagine synthesis and hydrolysis are components of the alanine, aspartate, and glutamate (Ala-Asp-Glu) metabolism pathway.</p>
<p>For the bottom-up analyses, significance criteria were loosened to include genes DE in a minimum of 15 of the 30 DE datasets, rather than all 30 datasets as used in the top-down analysis. This was to expand the DE data to include genes outside of the top-down lists. Significance criteria were maintained at a log<sub>2</sub>FC of at least 1.5 and p-value &lt;0.01. DE data were searched for any DE genes with annotations including the terms &#x201c;alanine&#x201d;, &#x201c;aspartate&#x201d;, &#x201c;glutamate&#x201d;, &#x201c;asparagine&#x201d;, and &#x201c;oxaloacetate&#x201d;, which fit these significance criteria. <xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;2</bold>
</xref> summarizes the log<sub>2</sub>FC in expression for all significantly DE genes with these annotations. <xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5</bold>
</xref> shows the relative expression data across all samples in this study as a heatmap; Pearson&#x2019;s coefficient relationship between genes in each list was used to organize the heatmap. These heatmaps provide a visual summary of the expression data and the relationships between the genes, while the information in <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref> provides the specific log<sub>2</sub>FC DE data (|log<sub>2</sub>FC| &gt; 1.5).</p>
<fig id="f5" position="float">
<label>Figure&#xa0;5</label>
<caption>
<p>Heatmap of bottom-up soybean genes with ontologies related to <bold>(A)</bold> asparagine, <bold>(B)</bold> alanine, <bold>(C)</bold> aspartate, <bold>(D)</bold> glutamate, and <bold>(E)</bold> oxaloacetate. Heatmaps were created using Heatmapper (<xref ref-type="bibr" rid="B5">Babicki et&#xa0;al., 2016</xref>). On the left side of each heatmap are the relationships between the genes. On the right of each heatmap is the corresponding gene name. Replicate sample names are given at the bottom of the map. Clustering was calculated using average linkage, and distance measurements were calculated using Pearson&#x2019;s coefficient. The row z-score for each heatmap is given; blue represents lower expression, and red represents higher expression. KEGG, Kyoto Encyclopedia of Genes and Genomes; DE, differential expression.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-14-1260393-g005.tif"/>
</fig>
<sec id="s3_3_1">
<label>3.3.1</label>
<title>Asparagine-related genes</title>
<p>A total of nine unique asparagine-related gene IDs were identified (based on criteria of p-value &lt;0.01 and a log<sub>2</sub> fold change of 1.5 in a minimum 15 of 30 datasets) to be DE between East and West across all 30 datasets with 189 total instances of DE (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). Seven genes (<italic>Glyma.07G257700</italic>, <italic>Glyma.13G279200</italic>, <italic>Glyma.15G073100</italic>, <italic>Glyma.15G105100</italic>, <italic>Glyma.11G171400</italic>, <italic>Glyma.11G170300</italic>, and <italic>Glyma.18G061100</italic>) with asparagine synthetase (AS) annotations were found to be downregulated across all lines in all three western locations (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5A</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). <italic>Glyma.11G171400</italic>, <italic>Glyma.11G170300</italic>, and <italic>Glyma.18G061100</italic> were all identified as AS (E.C.6.3.5.4) in <italic>G. max</italic>; these three genes were downregulated in Brandon and Morden, but DE in Saskatoon was limited to three instances, two of which were upregulated. Additionally, <italic>Glyma.13G279200</italic> and <italic>Glyma.15G105100</italic> had more instances of DE in Brandon and Morden than in Saskatoon. <italic>Glyma.13G279200</italic> encodes a stem-specific protein TSJT1, and <italic>Glyma.15G105100</italic> encodes a Wali7 domain-containing protein in <italic>G. max</italic>; both genes have PANTHER annotations of AS and TAIR10 identified the top <italic>Arabidopsis</italic> homolog is an aluminum-induced protein with YGL and LRDR motifs (AILP1). <italic>Glyma.07G257700</italic> was mostly found to be downregulated across the Saskatoon DE datasets (10), but also found to be downregulated in Brandon (2) and Morden (3); this gene had an NCBI identity of stem-specific protein TSJT1 in <italic>G. max</italic> based on model evidence. <italic>Glyma.15G073100</italic> was downregulated across all 30 DE datasets; this gene also encodes a stem-specific protein TSJT1 in <italic>G. max</italic> (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>).</p>
<p>Two genes (<italic>Glyma.17G238200</italic> and <italic>Glyma.04G042100</italic>) identified as asparaginase (ASPG) (E.C.3.5.1.1) in <italic>G. max</italic> were found to be largely upregulated across all 30 datasets, with one (<italic>Glyma.17G238200</italic>, <sc>l</sc>-asparaginase) persistently upregulated across all 30 datasets and the other (<italic>Glyma.04G042100</italic>, asparaginase 2) upregulated across 19 datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5A</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>).</p>
<p>The asparagine-gene IDs were converted to NCBI IDs and run through KEGG pathway mapping software. Four genes with AS annotations (PANTHER) did not map to any KEGG pathway data, including the three identified as TSJT1 (<italic>Glyma.07G257700</italic>, <italic>Glyma.13G279200</italic>, and <italic>Glyma.15G073100</italic>) and the gene encoding a Wali7 domain-containing protein in <italic>G. max</italic> (<italic>Glyma.15G105100</italic>) (<xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). Three genes with AS annotations (<italic>Glyma.11G171400</italic>, <italic>Glyma.11G170300</italic>, and <italic>Glyma.18G061100</italic>) mapped to E.C.6.3.5.4, and two genes with ASPG annotations (<italic>Glyma.17G238200</italic> and <italic>Glyma.04G042100</italic>) mapped to E.C.3.5.1.1 on the alanine, aspartate, and glutamate (Ala-Asp-Glu) metabolism pathway in soybeans (gmx00250) (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). The same five genes were also mapped to the biosynthesis of the secondary metabolites pathway (gmx01110) and the full metabolic pathway known for <italic>G. max</italic> (gmx01100). The three known <italic>G. max</italic> AS genes that mapped <italic>via</italic> KEGG to gmx00250 also mapped to the biosynthesis of the amino acid pathway (gmx01230); the two ASPG genes that mapped to gmx00250 also mapped to the cyanoamino acid metabolism pathway (gmx00460) (<xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>).</p>
<fig id="f6" position="float">
<label>Figure&#xa0;6</label>
<caption>
<p>Differential expression of soybean genes represented on the KEGG alanine, aspartate, and glutamate metabolism pathway map (gmx00250). Numbers in boxes are Enzyme Commission numbers comprising one or more proteins. Green-colored boxes represent enzyme sets identified in soybeans (<italic>Glycine max</italic>); white boxes are not known in soybeans. Red-bordered boxes represent enzyme sets encoded by genes that are upregulated in western locations relative to east, and blue-bordered boxes represent downregulated genes. The half-blue-half-red box (2.6.1.44) was found to be upregulated for serine&#x2014;glyoxylate and downregulated for alanine&#x2014;glyoxylate aminotransferase; they fall under the same E.C. KEGG, Kyoto Encyclopedia of Genes and Genomes.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-14-1260393-g006.tif"/>
</fig>
<p>Using DNASTAR MegAlign Pro (v17.4.3), we aligned the following: the amino acid sequences for <italic>Glyma.07G257700</italic>, <italic>Glyma.13G279200</italic>, <italic>Glyma.15G073100</italic>, <italic>Glyma.15G105100</italic>, <italic>Glyma.11G171400</italic>, <italic>Glyma.11G170300</italic>, and <italic>Glyma.18G061100</italic>; AS full sequence in <italic>G. max</italic> (XP_003538618.1; 566 amino acids); and the AS <italic>G. max</italic> glutamine amidotransferase (GATase) domain type-2 (amino acids 2&#x2013;185). <xref ref-type="supplementary-material" rid="SF2">
<bold>Supplementary Figure&#xa0;2</bold>
</xref> shows the protein sequence alignment of the AS sequences DE in this study and the known AS amino acid sequence, including the functional domain GATase. Using this analysis, we identified that the TJST1 protein sequences lack the Cys at amino acid 2 in the protein sequence, which is the active site for GATase activity (<xref ref-type="supplementary-material" rid="SF2">
<bold>Supplementary Figure&#xa0;2</bold>
</xref>). In the Wali7 domain-containing protein sequence encoded by <italic>Glyma.15G105100</italic>, there is Glu instead of the Cys at the GATase active site. Further, the binding site for <sc>l</sc>-glutamine (amino acid position 98) is an aspartate in the AS and the GATase domain sequences, but in the TJST1 and Wali7 domain-containing proteins, glutamate is encoded (<xref ref-type="supplementary-material" rid="SF2">
<bold>Supplementary Figure&#xa0;2</bold>
</xref>).</p>
</sec>
<sec id="s3_3_2">
<label>3.3.2</label>
<title>Alanine-related genes</title>
<p>Seven genes were found to be DE between East and West, six of which were downregulated (<italic>Glyma.01G129400</italic>, <italic>Glyma.03G040600</italic>, <italic>Glyma.07G247900</italic>, <italic>Glyma.11G006500</italic>, <italic>Glyma.17G026200</italic>, and <italic>Glyma.18G021000</italic>), and one (<italic>Glyma.18G116900</italic>) was upregulated in the West (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). <italic>Glyma.01G129400</italic> was the most persistently downregulated gene, with 28 instances of downregulation out of the 30 DE datasets; this gene is predicted to be an alanine&#x2013;glyoxylate aminotransferase 2 homolog 2 (mitochondrial) (E.C.2.6.1.44; <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>) in <italic>G. max</italic> and BLASTP identified alanine glyoxylate aminotransferase-like protein in <italic>Medicago truncatula</italic> as the most closely related protein (<xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). <italic>Glyma.18G021000</italic> was downregulated in the West in 26 of 30 datasets; this gene is uncharacterized in <italic>G. max</italic> but most closely related to an alanine&#x2013;glyoxylate aminotransferase 2, (mitochondrial, fragment) in <italic>Tupaia chinensis</italic> (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>).</p>
<p>
<italic>Glyma.18G116900</italic> was the only gene upregulated in the West with alanine-related annotations that fit our stringent criteria (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). This enzyme falls within the same pathway enzymatic element (E.C.2.6.1.44) but encodes a serine&#x2013;glyoxylate aminotransferase-like protein in <italic>G. max</italic> (<xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). Two of the downregulated genes mapped to the Ala-Asp-Glu metabolism pathway <italic>via</italic> KEGG mapping: <italic>Glyma.01G129400</italic> and <italic>Glyma.03G040600</italic>. These two genes both encode alanine&#x2013;glyoxylate aminotransferase 2 homolog 2 and homolog 3 and are also enzymatic components mapping to E.C.2.6.1.44 (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;3</bold>
</xref>).</p>
</sec>
<sec id="s3_3_3">
<label>3.3.3</label>
<title>Aspartate-related genes</title>
<p>In total, seven genes with &#x201c;aspartate&#x201d; in their annotations were found to be DE between East and West (p-value &lt;0.01, |log<sub>2</sub>FC| &#x2265; 1.5, minimum 15 of 30 datasets) with a total of 163 instances of significant DE across all 30 datasets (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>).</p>
<p>Four genes were found to be upregulated (<italic>Glyma.04G042100</italic>, <italic>Glyma.05G007600</italic>, <italic>Glyma.17G238200</italic>, and <italic>Glyma.19G159300</italic>); two of these genes encode ASPGs (E.C.3.5.1.1; <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>) and were present on the upregulated asparagine-related gene list (<italic>Glyma.04G042100</italic> and <italic>Glyma.17G238200</italic>) (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). <italic>Glyma.05G007600</italic> was upregulated in all 30 datasets; <italic>Glyma.05G007600</italic> encodes <sc>l</sc>-aspartate oxidase (E.C.1.4.3.16; <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>) in <italic>G. max</italic>. <italic>Glyma.19G159300</italic> was upregulated in 25 of 30 datasets; this gene encodes a lifeguard 4 protein in <italic>G. max</italic> and is most closely related to the gene encoding the glutamate-binding protein in <italic>Arabidopsis thaliana</italic> (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>).</p>
<p>Three genes were largely downregulated in the West (<italic>Glyma.02G015800</italic>, <italic>Glyma.14G111800</italic>, and <italic>Glyma.17G116500</italic>; <xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5C</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). <italic>Glyma.14G111800</italic> encodes an aspartate aminotransferase P2 (E.C.2.6.1.1) and was found to be downregulated in 15 of 30 DE datasets; however, 10 of these instances were in Saskatoon, and two and three instances were in Brandon and Morden, respectively. <italic>Glyma.02G015800</italic> was found to be downregulated in 16 of 30 datasets; this gene encodes fumatate hydratase 1 (E.C.4.2.1.2) in <italic>G. max</italic>. <italic>Glyma.17G116500</italic> was found to be downregulated across 28 of 30 datasets; this gene encodes broad specificity amino-acid racemase RacX in <italic>G. max</italic> and is homologous to aspartate-glutamate racemase family proteins in <italic>Populus trichocarpa</italic> and <italic>A. thaliana</italic> (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>).</p>
</sec>
<sec id="s3_3_4">
<label>3.3.4</label>
<title>Glutamate-related genes</title>
<p>Thirty genes with glutamate-inclusive annotations were found to be DE between eastern- and western-grown soybeans, which was made up of 13 upregulated genes and 17 downregulated genes (|log2FC| &#x2265; 1.5, p-value &lt;0.01, in at least 15 of 30 DE datasets) (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). Of the upregulated glutamate-related genes, two were found to be upregulated across all 30 East <italic>vs.</italic> West DE datasets, both of which encode proline dehydrogenase <italic>Glyma.13G049700</italic> (proline dehydrogenase 2, mitochondrial) and <italic>Glyma.19G043000</italic> (proline dehydrogenase). <italic>Glyma.19G111000</italic> encodes a glutamate dehydrogenase 1-like protein and was the only gene among upregulated glutamate-related genes to map to an E.C. using KEGG:glutamate dehydrogenase (E.C.1.4.1.3) (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). <italic>Glyma.13G233000</italic>, which encodes glutamate receptor 2.7, was downregulated in the West across all 30 DE datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). A number of other genes were found to be downregulated across nearly all western-grown soybeans, including <italic>Glyma.08G129600</italic> (cationic amino acid transporter 1) (29), <italic>Glyma.17G116500</italic> (broad specificity amino-acid racemase RacX) (28), and <italic>Glyma.03G069400</italic> (&#x3b4;-1-pyrroline-5-carboxylate synthase; ALDH18B3) (28) (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). Four downregulated genes mapped to the Ala-Asp-Glu metabolism pathway (gmx00250): <italic>Glyma.01G099800</italic>, <italic>Glyma.04G236900</italic>, <italic>Glyma.09G168900</italic>, and <italic>Glyma.18G041100</italic>. <italic>Glyma.01G099800</italic>, another &#x3b4;-1-pyrroline-5-carboxylate synthase (ALDH18B1), was downregulated in 20 of 30 datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>); this gene mapped to E.C.1.2.1.88 class of oxidoreductases in the production of glutamate (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). <italic>Glyma.04G236900</italic> encodes a NADH-dependent glutamate synthase and was found to be downregulated in 15 datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>) and mapped to the glutamate synthase enzymatic component (E.C.1.4.1.14) in the Ala-Asp-Glu pathway (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). <italic>Glyma.09G168900</italic> encodes a glutamate decarboxylase and was downregulated in 27 of 30 datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>); KEGG mapping identified this gene as an enzyme component included in the Ala-Asp-Glu metabolism pathway, glutamate decarboxylase (E.C.4.1.1.15) (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). <italic>Glyma.18G041100</italic> is a known glutamine synthetase nodule isozyme and was found to be downregulated in 23 of 30 datasets; this gene mapped to the glutamine synthetase enzyme component (E.C.6.3.1.2) (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). For the full annotated list of DE glutamate-related genes and relative expression, see <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>.</p>
</sec>
<sec id="s3_3_5">
<label>3.3.5</label>
<title>Oxaloacetate-related genes</title>
<p>A total of eight genes were identified to be DE between East and West, seven of which were upregulated (<italic>Glyma.04G086300</italic>, <italic>Glyma.06G087800</italic>, <italic>Glyma.08G201200</italic>, <italic>Glyma.13G354900</italic>, <italic>Glyma.15G019300</italic>, <italic>Glyma.15G055600</italic>, and <italic>Glyma.18G261800</italic>) and one of which was downregulated (<italic>Glyma.01G008500</italic>) (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5E</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). From model-based data, NCBI identities of six of the oxaloacetate-related genes are NADP-dependent malic enzymes in <italic>G. max</italic>, five of which are upregulated (<italic>Glyma.04G086300</italic>, <italic>Glyma.06G087800</italic>, <italic>Glyma.08G201200</italic>, <italic>Glyma.13G354900</italic>, and <italic>Glyma.15G019300</italic>) and a single downregulated gene (<italic>Glyma.01G008500</italic>) (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5E</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). One of the upregulated genes (<italic>Glyma.15G055600</italic>) is uncharacterized in <italic>G. max</italic>; however, BLASTP identified the most closely related protein to be a 2-oxoglutarate/malate translocator in <italic>M. truncatula</italic>. All six known malic enzyme genes (five upregulated and one downregulated) mapped to malate dehydrogenase (E.C.1.1.1.40), a component of the pyruvate metabolism pathway (gmx00620) and the carbon fixation in photosynthetic organism pathway (gmx00710) (map data not shown; see <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). <italic>Glyma.13G354900</italic> (malic enzyme) was upregulated in the West across all 30 DE datasets, and four genes were found to be DE in nearly all 30 datasets: <italic>Glyma.15G019300</italic> was upregulated in 29 of 30 datasets, <italic>Glyma.18G261800</italic> was upregulated in 28 of 30 datasets, <italic>Glyma.06G087800</italic> was upregulated in 27 of 30 datasets, and <italic>Glyma.15G055600</italic> was upregulated in 26 of 30 datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5E</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>).</p>
</sec>
<sec id="s3_3_6">
<label>3.3.6</label>
<title>QTL analysis</title>
<p>QTL analysis was performed using the bottom-up gene lists to determine if any of the DE genes of interest fall within seed protein or oil QTLs. <xref ref-type="fig" rid="f7">
<bold>Figures&#xa0;7A&#x2013;T</bold>
</xref> depicts <italic>G. max</italic> chromosomes 1&#x2013;20 with known seed protein and oil QTLs mapped alongside the bottom-up genes of interest. Genes that are found within large-spanning QTL regions are denoted with an asterisk. <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref> provides the map details for all QTLs and genes depicted in <xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>; the green highlight in this table indicates regions that fall within a major spanning QTL. On chromosome 11, <italic>Glyma.11G170300</italic> (AS; 18,242,402 cM) and <italic>Glyma.11G171400</italic> (AS; 18,518,857 cM) fall within a large oil QTL, seed linoleic 5-g4 (10,969,418&#x2013;25,595,388 cM) (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7K</bold>
</xref>). On chromosome 15, <italic>Glyma.15G055600</italic> (2-oxoglutarate/malate translocator-like protein; 4,350,474 cM) and the most persistently downregulated asparagine-related gene <italic>Glyma.15G073100</italic> (stem-specific protein TSJT1; 5,604,443 cM) fall within a large oil QTL, seed oil 11-g5 (4,148,354&#x2013;5,633,343 cM) (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7O</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>). <italic>Glyma.05G007600</italic> (<sc>l</sc>-aspartate oxidase) is in close proximity to many oil QTLs on chromosome 5 (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7E</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>). Just outside of the QTL seed oil 11-g5 is <italic>Glyma.15G105100</italic> (stem-specific protein TSJT1; 5,604,443 cM), among the downregulated asparagine-related genes (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7O</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>). Within a large oil QTL on chromosome 16 called seed palmitic 6-g4 (3,002,525&#x2013;4,148,354 cM) includes <italic>Glyma.16G038300</italic> (methionine synthase; 3,617,280 cM) (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7P</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>); this gene was found in the glutamate-related data to be upregulated in 26 of 30 East <italic>vs.</italic> West DE datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). Chromosome 17 has a large protein QTL, seed protein 9-g5 (38,930,849&#x2013;40,629,216 cM), which includes one of the most persistently DE ASPG genes, <italic>Glyma.17G238200</italic> (ASPG; 39,357,077 cM). Two genes on chromosome 19, <italic>Glyma.19G042900</italic> (proline dehydrogenase 2; 6,251,912 cM) and <italic>Glyma.19G043000</italic> (proline dehydrogenase; 6,272,211 cM), fall within a very large seed protein QTL, seed protein 9-g6 (2,437,848&#x2013;8,172,484 cM); these two genes were upregulated in 20 and 30 datasets, which of course makes <italic>Glyma.19G042900</italic> among the most persistently DE genes in the glutamate analysis.</p>
<fig id="f7" position="float">
<label>Figure&#xa0;7</label>
<caption>
<p>QTL maps for all 20 <italic>Glycine max</italic> chromosomes including all known protein and oil QTL information. Chromosome number is written above each map. Larger-spanning QTL regions are represented as bars to the right of each map; single-point QTLs (SNPs) are mapped directly onto the chromosome. Protein QTLs are in dark green; lipid QTLs are in mustard; asparagine-related genes are in red; alanine-related genes are in orange-red; oxaloacetate-related genes are in pink; aspartate-related genes are in royal blue; glutamate-related genes are in lime green; gene positions in cM are included in black text. Genes with an asterisk fall within spanning QTL regions. QTL, quantitative trait locus; SNPs, single-nucleotide polymorphisms. Chromosomes are in order from 1&#x2013;20: <bold>(A)</bold> chromosome 1; <bold>(B)</bold> chromosome 2; <bold>(C)</bold> chromosome 3; <bold>(D)</bold> chromosome 4; <bold>(E)</bold> chromosome 5; <bold>(F)</bold> chromosome 6; <bold>(G)</bold> chromosome 7; <bold>(H)</bold> chromosome 8; <bold>(I)</bold> chromosome 9; <bold>(J)</bold> chromosome 10; <bold>(K)</bold> chromosome 11; <bold>(L)</bold> chromosome 12; <bold>(M)</bold> chromosome 13; <bold>(N)</bold> chromosome 14; <bold>(O)</bold> chromosome 15; <bold>(P)</bold> chromosome 16; <bold>(Q)</bold> chromosome 17; <bold>(R)</bold> chromosome 18; <bold>(S)</bold> chromosome 19; <bold>(T)</bold> chromosome 20.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-14-1260393-g007.tif"/>
</fig>
</sec>
</sec>
</sec>
<sec id="s4" sec-type="discussion">
<label>4</label>
<title>Discussion</title>
<p>In this research, we investigated differences in the expression of genes between soybeans grown in East and West Canada in an effort to uncover DE genes and pathways that may contribute to the difference in seed protein content observed between the two locations over the past two decades (<xref ref-type="bibr" rid="B10">Canadian Grain Commission, 2022</xref>). Ten soybean genotypes were compared between East (Ottawa ON) and three different western locations (Morden MB, Brandon MB, and Saskatoon SK) in order to relieve genotypic and location biases for this large-scale RNA-seq and DE analysis. In this study, top-down (holistic) and bottom-up (keyword annotation search) approaches were used for the analysis of DE data to investigate the genes that are most consistently DE between East and West with putative roles influencing seed protein biosynthesis and accumulation.</p>
<sec id="s4_1">
<label>4.1</label>
<title>East <italic>vs.</italic> West transcriptome variability</title>
<p>The PCA plots in <xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref> show a clear separation between variability of expression in the East from variability of expression across all three western locations, implicating that all lines across the western-grown soybeans are behaving similarly to the others from the West and all eastern-grown soybeans are behaving similarly to others in the East. These plots were constructed using RNA-seq variability as the principal components, and with these results, we observe a clear difference in expression variability between the two geographic areas (East and West). Considering the fact that RNA-seq variability was used to assess identical genotypes grown in four different environments, it can be confidently concluded that across all 10 lines, soybeans in the West show differential transcriptomics than eastern-grown counterparts. Each plot in <xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref> represents a different genotype; colored data points represent replicates from each location, and gray data points correspond with the data on all other plots (all other genotypes). It is evident that all samples from the West are behaving more or less the same across all 10 genotypes, and the same can be said for the East. Indeed, many factors, both genetic and environmental, cumulatively influence resultant seed protein content (<xref ref-type="bibr" rid="B63">Wang et&#xa0;al., 2019</xref>); thus, we used big-data ontology, pathway, and QTL analyses to refine genes that were consistently DE between East and West across 30 individual datasets. The clustering of the transcriptome data from the three West locations is clearly separate from the data from the East across all 10 genotypes, indicating that location biases are minimized.</p>
</sec>
<sec id="s4_2">
<label>4.2</label>
<title>Amino acid biosynthesis and seed protein, with a focus on the Ala-Asp-Glu pathway</title>
<p>Among enriched pathways, biosynthesis of secondary metabolites (gmx01110) was prominently DE between samples grown in the two regions. GO and KEGG analyses identified different amino acid biosynthetic pathways were differently regulated between East and West, including aromatic amino acids (GO:0009072; gmx00400), sulfur-containing amino acids (GO:0000096; gmx00270), glycine (GO:0006546), and glutamate aspartate and asparagine (GO:0006537; GO:0033345; GO:0008734; gmx00250). A focus on the Ala-Asp-Glu biosynthesis pathway (gmx00250) was chosen because of the relationship between nitrogen assimilates, asparagine, and seed protein at maturity. Previous investigations into the DE between vegetable soybean and grain soybean found that the Ala-Asp-Glu metabolic pathway was highly enriched, as well as fatty acid biosynthesis and metabolism, carbon (starch and sucrose) metabolism and transport, arginine and proline metabolism, and glycolysis/gluconeogenesis, all of which influence the attributes (including protein) of the resulting seed (<xref ref-type="bibr" rid="B11">Chen et&#xa0;al., 2022</xref>).</p>
<p>During embryo development in plants, sucrose provides a source of carbon, and glutamine and asparagine are the main nitrogenous assimilate sources (<xref ref-type="bibr" rid="B49">Rainbird et&#xa0;al., 1984</xref>). Asparagine plays a key role in the source (root nodules)&#x2013;sink (seeds, mainly) translocation relationship (<xref ref-type="bibr" rid="B33">Lam et&#xa0;al., 1996</xref>). Asparagine has a relatively high nitrogen:carbon ratio and is biochemically stable, which make it ideal for nitrogen transport and storage. A careful balance exists between asparagine biosynthesis and degradation to maintain asparagine concentration. Asparagine represents up to 50% of the total free amino acids in the developing cotyledon (<xref ref-type="bibr" rid="B29">Hern&#xe1;ndez-Sebasti&#xe0; et&#xa0;al., 2005</xref>). AS is the major enzyme responsible for synthesizing asparagine. Typically, two or more AS genes are found in plants. High AS activity in the cotyledons of the germinating seed, as well as in mature root nodules, supports the idea that asparagine acts as a nitrogen transport system in legume plants (<xref ref-type="bibr" rid="B33">Lam et&#xa0;al., 1996</xref>). AS generates asparagine from aspartate by using glutamine or ammonia as a substrate for the transfer of the amide group to aspartic acid in an ATP-dependent reaction catalyzed by magnesium (<xref ref-type="bibr" rid="B37">Lea et&#xa0;al., 2007</xref>). AS proteins are categorized as either AS-A (a.k.a. AsnA) or AS-B (a.k.a. AsnB); AS-B family proteins, found in both prokaryotes and eukaryotes, can use both ammonia and glutamine as a nitrogen donor but prefer glutamine (<xref ref-type="bibr" rid="B41">Manhas et&#xa0;al., 2014</xref>). Glutamine-dependent AS is the main asparagine biosynthesis pathway in plants (<xref ref-type="bibr" rid="B33">Lam et&#xa0;al., 1996</xref>). ASPGs are ubiquitous across all domains of life. ASPG breaks down the isoaspartyl peptide bond in asparagine to aspartate and ammonia, which are then reassimilated through the glutamine synthase/glutamate synthase cycle (<xref ref-type="bibr" rid="B24">Gomes and Sodek, 1984</xref>; <xref ref-type="bibr" rid="B26">Haga and Sodek, 1987</xref>). In soybeans, ASPG activity is directly associated with a reduction in free asparagine by up to 18% while simultaneously increasing the amount of aspartate by up to 60% (<xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>). ASPG activity was also associated with a reduction in total nitrogen by 9%&#x2013;13% and an increased concentration of seed oil by 5%&#x2013;8% in soybeans (<xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>).</p>
<p>From this study, we observe differences in the expression of genes related to asparagine metabolism between soybeans grown in eastern and western Canada. The western-grown soybeans show downregulation of AS compared to eastern-grown soybeans (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5A</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). This indicates that soybeans in the West are not synthesizing asparagine to the same degree as soybeans in the East, which may be directly attributed to the seed protein content at maturity. The western soybeans showed consistent upregulated expression of ASPG compared to the eastern counterparts. One of the most persistently upregulated ASPG genes, <italic>Glyma.17G238200</italic>, falls within a large protein QTL, seed protein 9-g5 (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>). Further, the most persistently downregulated asparagine-related gene, <italic>Glyma.15G073100</italic> (stem-specific protein TSJT1), is within the major oil QTL, seed oil 11-g5 on chromosome 15 (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>), another chromosome highly enriched for seed protein and oil QTLs. This might suggest that these genes are linked to seed protein and oil contents. AS and TSJT1 both have annotations that include &#x201c;asparagine synthetase&#x201d;; however, the sequence-based analysis uncovered a major difference in their protein sequences: TSJT1 protein sequences do not contain the GATase activity position 2 Cys, and the Wali7 domain-containing protein sequence has a Glu at amino acid position 2 (<xref ref-type="supplementary-material" rid="SF2">
<bold>Supplementary Figure&#xa0;2</bold>
</xref>). AS genes <italic>Glyma.11G171400</italic>, <italic>Glyma.11G170300</italic>, and <italic>Glyma.18G061100</italic> are extremely downregulated in Brandon and Morden, but not found to be downregulated in Saskatoon (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). One TSJT1-encoding gene, <italic>Glyma.15G073100</italic>, was downregulated across all 30 datasets, while the other TSJT1-encoding genes were found to be downregulated mostly in Brandon and Morden (<italic>Glyma.13G279200</italic>) or Saskatoon (<italic>Glyma.07G257700</italic>).</p>
<p>The concentration of protein in mature soybeans is strongly associated with free asparagine in the plant during development, making it an ideal pathway for further investigation (<xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>). With the observations made in this study that ASPG is highly upregulated in the West and AS is highly downregulated in the West, it is entirely plausible that differences in asparagine metabolism are influencing the seed protein accumulation difference between eastern- and western-grown soybeans. It would be of interest to soybean breeding programs to consider increasing AS expression as an engineering target when designing high-protein soybean lines. AS1 overexpression in <italic>A. thaliana</italic> resulted in increased free asparagine levels and increased seed protein concentration (<xref ref-type="bibr" rid="B34">Lam et&#xa0;al., 2003</xref>). An increase in AS1 expression in soybean leaves had a positive correlation to seed protein concentration (<xref ref-type="bibr" rid="B60">Wan et&#xa0;al., 2006</xref>). In soybean roots, increased AS1 expression was correlated with an increased ratio of asparagine:aspartate in xylem sap headed to shoots, implicating more asparagine being transported to aerial tissues (<xref ref-type="bibr" rid="B4">Antunes et&#xa0;al., 2008</xref>).</p>
<p>Increased expression of asparagine aminohydrolases in western soybeans indicates that these plants are breaking down available asparagine to recycle the components, most specifically the nitrogen. When nitrogen is limited, hydrolyzing asparagine provides a source of nitrogen to be redirected into other processes. Because of the central intermediary relationship between alanine and/or serine and asparagine transamination (both amino acids can act as a substrate), increased ASPG expression, which leads to a reduction in freely available asparagine (<xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>), could logically be associated with equally proportionate increases in serine and alanine. Interestingly, in the data presented in this study, an increase in the gene encoding a serine&#x2013;glyoxylate aminotransferase 2 (<italic>Glyma.18G116900</italic>) was observed, while multiple genes encoding alanine&#x2013;glyoxylate aminotransferases (<italic>Glyma.01G129400</italic>, <italic>Glyma.03G040600</italic>, and <italic>Glyma.18G021000</italic>) were highly downregulated (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). This means that different activities between two different enzymes of the same enzyme component (E.C.2.6.1.44) are simultaneously ongoing in western-grown soybeans, as presented in <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref> by the half-blue-half-red box.</p>
<p>The asparagine synthesis pathway within the Ala-Asp-Glu metabolism pathway shows two other enzyme components that directly influence asparagine biosynthesis/metabolism: E.C.6.3.1.1 and E.C.3.5.1.38 (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>). E.C.6.3.1.1 is an AS-A family AS that is an aspartate&#x2013;ammonia ligase and was not found to be DE within our data. This is expected because plant AS is of the AS-B family of AS proteins and primarily uses glutamate as a N donor source rather than aspartate (<xref ref-type="bibr" rid="B41">Manhas et&#xa0;al., 2014</xref>). The KEGG pathway depicted in <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref> KEGG shows enzyme components that are known to be in soybeans by highlighting respective boxes in green. E.C.3.5.1.38 appears to be prokaryotic in nature as indicated by available information on KEGG and BRENDA enzyme databases. It should be noted that as a result of the decrease in asparagine, western-grown soybeans could be compensating by increasing the expression of other nitrogen-rich amino acid (arginine and lysine) metabolizing enzymes (<xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>), which were not covered in our bottom-up analysis and might serve as an interesting area for further research.</p>
<p>Alanine is one of the central intermediates in amino acid metabolism and a substrate of asparagine transaminase, the enzyme responsible for transferring the &#x3b1;-amino group between asparagine and glycine, alanine, serine, and homoserine (<xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>; <xref ref-type="bibr" rid="B66">Zhang et&#xa0;al., 2013</xref>; <xref ref-type="bibr" rid="B23">Gaufichon et&#xa0;al., 2015</xref>). There is a central intermediary relationship between asparagine transamination and alanine and/or serine in that both amino acids can act as a substrate (<xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>). Overall, the results from the alanine investigation indicate the downregulation of alanine-related genes, with six downregulated genes and one upregulated gene common across at least 50% of the DE datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). Significant downregulation of <italic>Glyma.01G129400</italic>, <italic>Glyma.03G040600</italic>, and <italic>Glyma.18G021000</italic>, three alanine&#x2013;glyoxylate aminotransferases, was observed in the West (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). An increase in ASPG leads to a decrease in asparagine and potentially an increase in alanine, which may in part explain the downregulation of alanine-related genes as a whole (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>). In <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>, E.C.2.6.1.44 is both up- and downregulated in the West (half-blue-half-red box). Western-grown soybeans appear to be upregulating the alanine-related gene, <italic>Glyma.18G116900</italic>, encoding serine&#x2013;glyoxylate aminotransferase 2 (E.C.2.6.1.44), which is also involved in serine-pyruvate transaminase activity (GO:0004760). Simultaneously, these soybeans are downregulating the expression of two alanine&#x2013;glyoxylate aminotransferases (<italic>Glyma.01G129400</italic> and <italic>Glyma.03G040600</italic>) (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). In addition to the Ala-Asp-Glu pathway, enrichment for the cyanoamino acid metabolism pathway (gmx00460) is a result of the two ASPG genes (<italic>Glyma.17G238200</italic> and <italic>Glyma.04G042100</italic>) (<xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>). High-protein soybean genotypes were found to have higher amounts of free asparagine and alanine in developing embryos than in low-protein genotypes (<xref ref-type="bibr" rid="B29">Hern&#xe1;ndez-Sebasti&#xe0; et&#xa0;al., 2005</xref>). Further, freely available 3-cyanoalanine was found to be highly correlated with seed protein and/or oil in soybeans (<xref ref-type="bibr" rid="B63">Wang et&#xa0;al., 2019</xref>). In an investigation into the genetic shift in soybeans over 24 years, it was found that newer soybean cultivars had a decrease in seed protein, alanine, and serine (<xref ref-type="bibr" rid="B17">de Borja Reis et&#xa0;al., 2020</xref>). The relationship between alanine, serine, and asparagine in soybean seed protein accumulation remains elusive, and further investigations into the relationship between these amino acids and protein should be explored.</p>
<p>The aspartate-family amino acid sub-pathway functions as a regulatory metabolic link with the tricarboxylic acid (TCA) cycle, biologically significant under extreme stress conditions, which deplete cellular energy (<xref ref-type="bibr" rid="B21">Galili, 2011</xref>), making it an essential metabolite for plant growth and stress acclimation (<xref ref-type="bibr" rid="B27">Han et&#xa0;al., 2021</xref>). Aspartate-family amino acids (lysine, threonine, methionine, and isoleucine) are synthesized in plants using aspartate as a central amino acid (<xref ref-type="bibr" rid="B21">Galili, 2011</xref>). In this study, <italic>Glyma.05G007600</italic> encoding <sc>l</sc>-aspartate oxidase (E.C.1.4.3.16) was significantly upregulated in the West. This gene is physically close on chromosome 5 to many known oil QTLs (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>). Chromosome 5 is highly enriched for QTL influencing protein and oil; however, the molecular mechanisms driven by these loci remain largely unknown (<xref ref-type="bibr" rid="B63">Wang et&#xa0;al., 2019</xref>). The closeness in proximity suggests that they are tightly linked, and following recombination, it would be advantageous for these genes to remain together in future progeny.</p>
<p>
<italic>Glyma.19G159300</italic> (lifeguard 4 protein in <italic>G. max</italic>) was upregulated in 25 of 30 datasets (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>); this gene is most closely related to the gene encoding the glutamate-binding protein in <italic>A. thaliana</italic> and is also closely related to the inhibitor of the apoptosis-promoting BAX1 protein in <italic>A. thaliana</italic> (<xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). Further investigations into the glutamate binding potential of this gene and the putative role it plays in signaling would be of merit to understanding the reason(s) for significant upregulation of <italic>Glyma.19G159300</italic>. It is likely that this protein plays a role in signaling and proliferation, potentially inhibiting apoptosis. Increasing expression of genes related to glutamate signaling and apoptosis suggests that western-grown soybeans are differently controlling cell proliferation compared to those in the East.</p>
<p>Glutamate has a remarkably wide range of biological roles because of the central position it plays in metabolism. It is suggested that glutamate compensates for the reduction in freely available asparagine by serving as a metabolite (nitrogen) storage molecule and behaves as an organic nitrogen signal in seedlings (<xref ref-type="bibr" rid="B25">Guti&#xe9;rrez et&#xa0;al., 2008</xref>; <xref ref-type="bibr" rid="B45">Pandurangan et&#xa0;al., 2012</xref>). This study uncovered 30 glutamate-related genes that are DE in at least 50% of the datasets. Within the glutamate data, genes for proline dehydrogenases (<italic>Glyma.13G049700</italic>, <italic>Glyma.19G042900</italic>, and <italic>Glyma.19G043000</italic>) were among the most upregulated in the West. Proline dehydrogenase catalyzes the oxidation of <sc>l</sc>-proline to &#x3b4;1-pyrroline-5-carboxylate, which provides a source of free electrons for transport (<xref ref-type="bibr" rid="B52">Servet et&#xa0;al., 2012</xref>). The two proline dehydrogenase genes on chromosome 19 both fall within a major seed protein QTL, seed protein 9-g6 (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>). Proline dehydrogenase catabolizes proline while simultaneously playing roles in energy, shuttling redox potential, and production of reactive oxygen species (ROS) to reach cellular homeostasis, adapt to the environment, and carry out physiological and pathological processes (<xref ref-type="bibr" rid="B52">Servet et&#xa0;al., 2012</xref>). A methionine synthase (<italic>Glyma.16G038300</italic>) was also highly upregulated in the West, which ultimately influences the sulfur-containing amino acid content of the developing seed. Methionine synthase is responsible for catalysis of 5-methyltetrahydropteroyltri-<sc>l</sc>-glutamate and <sc>l</sc>-homocysteine into <sc>l</sc>-methionine + tetrahydropteroyltri-<sc>l</sc>-glutamate in the synthesis of methionine and glutamate (<xref ref-type="bibr" rid="B64">Whitfield et&#xa0;al., 1970</xref>), a unique feature of some organisms that are able to convert Cys to Met under specific circumstances (<xref ref-type="bibr" rid="B9">Brosnan and Brosnan, 2006</xref>). Upregulation of this methionine synthase gene may suggest that western-grown soybeans are increasing glutamate in an attempt to compensate for a lack of free asparagine. If, coincidently, this also results in more Met production from <sc>l</sc>-homocysteine, the overall protein quality (in terms of 11S and 7S globulins) could be improved. Improved seed protein quality is an important consideration for soybean agriculture, particularly in regions where environmentally influenced decreases in seed protein levels are prominent (i.e., western Canada), and western-grown soybeans were found to have higher 11S:7S values than eastern-grown soybeans (<xref ref-type="bibr" rid="B15">Cober et&#xa0;al., 2023</xref>).</p>
<p>Also within the glutamate data are a number of DE glutamate receptor genes both upregulated (<italic>Glyma.07G226400</italic>, <italic>Glyma.10G213200</italic>, <italic>Glyma.13G049700</italic>, <italic>Glyma.13G093300</italic>, <italic>Glyma.13G172100</italic>, and <italic>Glyma.13G233400</italic>) and downregulated (<italic>Glyma.13G233000</italic>, <italic>Glyma.13G233300</italic>, and <italic>Glyma.17G067200</italic>) in the West (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). This points to specific genes involved in DE glutamate signaling between eastern- and western-grown soybeans. The specific ligands for these receptors would be an interesting area for further research on these genes.</p>
<p>As previously stated, the Ala-Asp-Glu metabolism pathway coordinates a metabolic link to the TCA cycle, directly influencing energy production or depletion (<xref ref-type="bibr" rid="B21">Galili, 2011</xref>). Oxaloacetate is both a product and a beginning component of the TCA cycle; levels of oxaloacetate give an indication of the ongoing level of energy metabolism. The data uncovered in this study indicate major upregulation of oxaloacetate-related genes in western-grown soybeans; of eight significantly DE genes, seven were upregulated, and one was downregulated (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5E</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>). All of the DE genes related to oxaloacetate are malic enzymes (E.C.1.1.1.40, pyruvate metabolism pathway gmx00620; <xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>), with the exception of <italic>Glyma.15G055600</italic>. <italic>Glyma.15G055600</italic>, a 2-oxoglutarate/malate translocator-like protein, was found within the oxaloacetate-related gene list (<xref ref-type="supplementary-material" rid="SM3">
<bold>Supplementary Table&#xa0;3</bold>
</xref>) and the only oxaloacetate-related gene that does not map to any pathway using KEGG (<xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>). However, <italic>Glyma.15G055600</italic> falls within the same major oil QTL as one of the most persistently downregulated asparagine-related genes (<italic>Glyma.15G073100</italic>), seed oil 11-g5 on chromosome 15 (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>).</p>
<p>Malic enzyme is one of the key enzymes linked to fatty acyl chain biosynthesis. <sc>l</sc>-Malic acid was found to be highly correlated with protein and oil contents in soybeans (<xref ref-type="bibr" rid="B63">Wang et&#xa0;al., 2019</xref>). A significant amount of pyruvate, the precursor of acetyl-CoA synthesis for lipid biosynthesis, is produced as a result of malic enzyme activity in soybeans (<xref ref-type="bibr" rid="B1">Allen and Young, 2013</xref>). In western-grown soybeans in this study, malic enzyme activity was highly upregulated, which suggests that these plants are likely producing higher levels of acetyl-CoA for fatty acid production. Because of this, malic enzyme activity is almost certainly one of the regulatory mechanisms underlying the inverse relationship between seed protein and oil contents in soybeans. This mechanism serves are an optimal target for genetic engineering/control of the decision between lipid and protein biosynthesis. Increasing expression of malic enzyme would make a molecular conduit to directing nitrogen and carbon toward lipid biosynthesis and away from protein biosynthesis, serving as a molecular switch for the accumulation of major seed storage biomolecules (<xref ref-type="bibr" rid="B42">Morley et&#xa0;al., 2023</xref>).</p>
</sec>
<sec id="s4_3">
<label>4.3</label>
<title>Sulfur-containing amino acid biosynthesis and seed protein</title>
<p>The ontologies sulfur amino acid metabolic process (GO:0000096), sulfur compound biosynthetic process (GO:0044272), sulfur compound metabolic process (GO:0006790), iron-sulfur cluster binding (GO:0051536), and more were all enriched and overrepresented within the top-down GO analysis (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). Further, the sulfur-containing amino acid Cys and Met metabolism pathway (gmx00270) was enriched within the DE data. These observations that transcription of genes related to the biosynthesis of sulfur-containing amino acids (Cys and Met) are DE between East and West likely play a role in the differences in seed protein quality seen between western- and eastern-grown soybeans (<xref ref-type="bibr" rid="B15">Cober et&#xa0;al., 2023</xref>). Sulfur-containing amino acids are essential to the formation of 11S storage proteins (glycinins) in soybean. Sedimentation coefficients (0.5-M ionic strength) are used to categorize seed storage proteins into 2S, 7S, 11S, and 15S fractions, of which the 11S and 7S fractions account for the majority of seed storage proteins (40% and 30% of seed storage protein, respectively) (<xref ref-type="bibr" rid="B48">Peng et&#xa0;al., 1984</xref>; <xref ref-type="bibr" rid="B55">Tsukada et&#xa0;al., 1986</xref>). Cys and Met are limited resources in soybeans, and tight control of biosynthesis of these amino acids is advantageous to glycinin production in soybeans. Because Cys and Met are vital to glycinin biosynthesis, the genes influencing expression and accumulation of sulfur-containing amino acids very likely influence glycinin accumulation (nutrient reservoir activity) and therefore seed protein content.</p>
</sec>
<sec id="s4_4">
<label>4.4</label>
<title>Other influential pathways on seed protein</title>
<p>Enriched among the data were many other pathways of interest That almost certainly influence seed protein content, including fatty acid metabolic processes (fatty acid biosynthetic process GO:0006633; lipid transport GO:0006869; lipid binding GO:0008289; gmx01212), circadian rhythm (circadian rhythm GO:0007623; regulation of circadian rhythm GO:0042752; gmx04712), nutrient storage (nutrient reservoir activity GO:0045735), and carbohydrate metabolism (carbohydrate metabolic process GO:0005975; carbohydrate-binding GO:0030246; regulation of carbohydrate metabolic process GO:0006109) (<xref ref-type="fig" rid="f3">
<bold>Figures&#xa0;3</bold>
</xref>, <xref ref-type="fig" rid="f4">
<bold>4</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). Western locations in this study are also further North than the Eastern locations and as a result experience longer photoperiods; the circadian rhythm of soybeans in these locations is almost certainly going to exhibit differences. An insertion/deletion in <italic>Glyma.20G85100</italic>, a circadian clock gene, was found to nearly perfectly correspond with high/low protein alleles for a QTL on chromosome 20, cqSeed protein-003 (<xref ref-type="bibr" rid="B20">Fliege et&#xa0;al., 2022</xref>). While <italic>Glyma.20G85100</italic> was not DE in our data, genes with circadian rhythm among their ontologies were overrepresented in both upregulated (20) and downregulated genes (23) (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref>; <xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). Response to UV (GO:0009411) is upregulated in western-grown soybeans, likely in response to reduced cloud cover in the prairies. Ontologies for microtubule cytoskeleton organization (GO:0000226), microtubule motor activity (GO:0003777), cytokinesis by cell plate formation (GO:0000911), cell wall biogenesis (GO:0042546), cell cycle (GO:0007049), and spindle assembly (GO:0051225) are highly enriched in the upregulated genes in our study (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref>, <xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). Soybeans in western Canada were found to be significantly taller than eastern-grown soybeans (<xref ref-type="bibr" rid="B15">Cober et&#xa0;al., 2023</xref>), which is likely influenced by these genes, suggesting the plants are spending more energy increasing in height than producing/filling seeds. Genes related to water stress were also found to be DE between East and West (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref>); response to desiccation (GO:0009269), water transport (GO:0006833), response to water depravation (GO:0009415), and water channel activity (GO:0015250) are all likely the result of the lower precipitation and lower relative humidity in the West.</p>
<p>Protein and oil compete for space in the seed, which results in a pushing/pulling relationship between these seed storage macronutrients (<xref ref-type="bibr" rid="B8">Breene et&#xa0;al., 1988</xref>; <xref ref-type="bibr" rid="B13">Clemente and Cahoon, 2009</xref>). Previous investigations into seed protein, oil, and yield identified an influential QTL on chromosome 20 between Satt496 and Satt239 (<xref ref-type="bibr" rid="B12">Chung et&#xa0;al., 2003</xref>); however, there were no DE genes found between these markers in our data. The relationship between seed storage and metabolism is entangled by the upstream carbohydrate metabolism decision-making steps that lead to protein and/or oil biosynthesis (40% and 20%, respectively) while also maintaining a proportion (~35%) of the seed space for stored carbohydrates (<xref ref-type="bibr" rid="B38">Liu, 1997</xref>). Seed storage proteins (glycinins, vicilins, and cupins) have the ontology nutrient reservoir activity (GO:0045735). Nutrient reservoir activity (GO:0045735) was overrepresented in the top-down downregulated genes but underrepresented (though present) in the upregulated GO enrichment (<xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). Known genes for glycinins and &#x3b2;-conglycinins (<italic>Glyma.03G163500</italic> GY1, <italic>Glyma.03G163500</italic> GY2, <italic>Glyma.19G164900</italic> GY3, <italic>Glyma.10G037100</italic> GY4, <italic>Glyma.13G123500</italic> GY5, <italic>Glyma.10G246300</italic> CG-1 &#x3b1;&#x2032;1, <italic>Glyma.20G148400</italic> CG-2 &#x3b1;2, <italic>Glyma.20G148300</italic> CG-3 &#x3b1;1, <italic>Glyma.20G146200</italic> CG-4 &#x3b2;1, and <italic>Glyma.20G148200</italic> CG-4 &#x3b2;2) were not found to be significantly DE across a majority of lines but were found to be DE in some lines (<xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Table&#xa0;1</bold>
</xref>). &#x3b2;-Conglycinin genes are reported to only be expressed in seeds in early embryogenesis; transcription spikes at mid-maturation and decreases before dormancy (<xref ref-type="bibr" rid="B28">Harada et&#xa0;al., 1989</xref>). These genes are not expressed in cotyledons or at maturity; thus, it is reasonable that we do not see notable DE between East and West in the present data; leaf tissue at the R5 (seed filling) stage was used for RNA-seq. A similar investigation into RNA-seq analysis of soybean pod data would provide further information into the DE of glycinins and &#x3b2;-conglycinins. The findings in this study share similarities with a similar study conducted in China on high and low seed protein varieties, which found DE of genes involved in the biosynthesis of amino acids and secondary metabolites, carbon metabolism, lipid metabolism, phenylpropanoid biosynthesis, and plant hormone signal transduction (<xref ref-type="bibr" rid="B65">Xu et&#xa0;al., 2022</xref>), although we examined DE of individual varieties across geographies and not DE between varieties.</p>
</sec>
</sec>
<sec id="s5" sec-type="conclusions">
<label>5</label>
<title>Conclusions</title>
<p>In this work, we identified genes persistently DE in 10 soybean varieties grown in three different locations in western Canada compared to an eastern Canada location. We pinpoint genes within specific metabolic processes that are likely key players in reduced protein content observed in western-grown soybeans, most pertinently genes encoding AS and ASPG. By investigating the differences in the expression of genes underlying nitrogen assimilation during seed development in soybeans grown in East and West Canada, we offer valuable information on the impact of geographic location on this pathway as well as potential avenues for breeding improvement opportunities. Further investigations into the lipid biosynthetic pathway, sulfur-containing amino acid biosynthesis pathway, and aromatic amino acid biosynthesis pathway would all likely provide key information into the differences in metabolic orchestration influenced by environment.</p>
</sec>
<sec id="s6" sec-type="data-availability">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/<xref ref-type="supplementary-material" rid="SM1">
<bold>Supplementary Material</bold>
</xref>. Further inquiries can be directed to the corresponding author.</p>
</sec>
<sec id="s7" sec-type="author-contributions">
<title>Author contributions</title>
<p>JH: Data curation, Formal Analysis, Methodology, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &amp; editing. MS: Conceptualization, Formal Analysis, Methodology, Writing &#x2013; review &amp; editing. GZ: Data curation, Formal Analysis, Software, Visualization, Writing &#x2013; review &amp; editing. MC: Data curation, Writing &#x2013; review &amp; editing. DL: Data curation, Methodology, Writing &#x2013; review &amp; editing. RM: Data curation, Writing &#x2013; review &amp; editing. KD: Data curation, Writing &#x2013; review &amp; editing. TW: Data curation, Resources, Supervision, Writing &#x2013; review &amp; editing. MH: Data curation, Writing &#x2013; review &amp; editing. BB: Data curation, Writing &#x2013; review &amp; editing. AH: Data curation, Resources, Writing &#x2013; review &amp; editing. FL: Data curation, Software, Writing &#x2013; review &amp; editing. AG: Formal Analysis, Supervision, Writing &#x2013; review &amp; editing. EC: Conceptualization, Formal Analysis, Funding acquisition, Methodology, Project administration, Resources, Writing &#x2013; review &amp; editing. BS: Conceptualization, Data curation, Formal Analysis, Funding acquisition, Investigation, Methodology, Project administration, Resources, Software, Supervision, Validation, Visualization, Writing &#x2013; review &amp; editing.</p>
</sec>
</body>
<back>
<sec id="s8" sec-type="funding-information">
<title>Funding</title>
<p>The authors declare financial support was received for the research, authorship, and/or publication of this article. This research was funded by Agriculture and Agri-Food Canada and the Canadian Field Crop Research Alliance (CFCRA).</p>
</sec>
<ack>
<title>Acknowledgments</title>
<p>We thank the field crew at the farms in Ottawa, Morden, Brandon, and Saskatoon. We would like to thank the Molecular Technology Lab at the Ottawa Research and Development Centre for their help with this project, with a special thank-you to Kasia Dadej. We would like to thank Agriculture and Agri-Food Canada (AAFC) and the Canadian Field Crop Research Alliance for their financial support. We would also like to thank G&#xe9;nome Qu&#xe9;bec (Montr&#xe9;al, Canada) for their contributions to RNA-sequencing. JH would like to thank MRW, SMB, HH, and MH.</p>
</ack>
<sec id="s9" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s10" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="s11" sec-type="supplementary-material">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fpls.2023.1260393/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fpls.2023.1260393/full#supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Table_1.xlsx" id="SM1" mimetype="application/vnd.openxmlformats-officedocument.spreadsheetml.sheet">
<label>Supplementary Table&#xa0;1</label>
<caption>
<p>Top-down analysis of up- and downregulated genes in each line-location analysis, and cumulatively DE genes. <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;2</bold>
</xref>: Bottom-up analysis summarizing asparagine, alanine, aspartate, glutamate, and oxaloacetate DE data. <xref ref-type="supplementary-material" rid="SM2">
<bold>Supplementary Table&#xa0;3</bold>
</xref>: KEGG pathway mapping information for top-down and bottom-up analyses. <xref ref-type="supplementary-material" rid="SM4">
<bold>Supplementary Table&#xa0;4</bold>
</xref>: MapChart chromosome map data for <italic>G. max</italic> chromosomes 1&#x2013;20. <xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Figure&#xa0;1</bold>
</xref>: High resolution image of top-down analysis KEGG pathway map (gmx01100). <xref ref-type="supplementary-material" rid="SF1">
<bold>Supplementary Figure&#xa0;2</bold>
</xref>: Amino acid sequence alignment of AS-related proteins.</p>
</caption>
</supplementary-material>
<supplementary-material xlink:href="Table_2.xlsx" id="SM2" mimetype="application/vnd.openxmlformats-officedocument.spreadsheetml.sheet"/>
<supplementary-material xlink:href="Table_3.xlsx" id="SM3" mimetype="application/vnd.openxmlformats-officedocument.spreadsheetml.sheet"/>
<supplementary-material xlink:href="Table_4.xlsx" id="SM4" mimetype="application/vnd.openxmlformats-officedocument.spreadsheetml.sheet"/>
<supplementary-material xlink:href="Image_1.jpeg" id="SF1" mimetype="image/jpeg"/>
<supplementary-material xlink:href="Image_2.jpeg" id="SF2" mimetype="image/jpeg"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Allen</surname> <given-names>D. K.</given-names>
</name>
<name>
<surname>Young</surname> <given-names>J. D.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>Carbon and nitrogen provisions alter the metabolic flux in developing soybean embryos</article-title>. <source>Plant Physiol.</source> <volume>161</volume>, <fpage>1458</fpage>&#x2013;<lpage>1475</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1104/pp.112.203299</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Anders</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Huber</surname> <given-names>W.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Differential expression analysis for sequence count data</article-title>. <source>Genome Biol.</source> <volume>11</volume>, <fpage>R106</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/gb-2010-11-10-r106</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Anders</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Pyl</surname> <given-names>P. T.</given-names>
</name>
<name>
<surname>Huber</surname> <given-names>W.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>HTSeq&#x2013;a Python framework to work with high-throughput sequencing data</article-title>. <source>Bioinformatics</source> <volume>31</volume>, <fpage>166</fpage>&#x2013;<lpage>169</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/bioinformatics/btu638</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Antunes</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Aguilar</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Pineda</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Sodek</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Nitrogen stress and the expression of asparagine synthetase in roots and nodules of soybean (Glycine max)</article-title>. <source>Physiol. Plant</source> <volume>133</volume>, <fpage>736</fpage>&#x2013;<lpage>743</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/j.1399-3054.2008.01092.x</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Babicki</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Arndt</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Marcu</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Liang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Grant</surname> <given-names>J. R.</given-names>
</name>
<name>
<surname>Maciejewski</surname> <given-names>A.</given-names>
</name>
<etal/>
</person-group>. (<year>2016</year>). <article-title>Heatmapper: web-enabled heat mapping for all</article-title>. <source>Nucleic Acids Res.</source> <volume>44</volume>, <fpage>W147</fpage>&#x2013;<lpage>W153</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/nar/gkw419</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bolger</surname> <given-names>A. M.</given-names>
</name>
<name>
<surname>Lohse</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Usadel</surname> <given-names>B.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Trimmomatic: a flexible trimmer for Illumina sequence data</article-title>. <source>Bioinformatics</source> <volume>30</volume>, <fpage>2114</fpage>&#x2013;<lpage>2120</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/bioinformatics/btu170</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bourgey</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Dali</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Eveleigh</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>K. C.</given-names>
</name>
<name>
<surname>Letourneau</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Fillon</surname> <given-names>J.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>GenPipes: an open-source framework for distributed and scalable genomic analyses</article-title>. <source>Gigascience</source> <volume>8</volume>, <elocation-id>giz037</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/gigascience/giz037</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Breene</surname> <given-names>W. M.</given-names>
</name>
<name>
<surname>Lin</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Hardman</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Orf</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>1988</year>). <article-title>Protein and oil content of soybeans from different geographic locations</article-title>. <source>J. Am. Oil Chem. Soc</source> <volume>65</volume>, <fpage>1927</fpage>&#x2013;<lpage>1931</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/BF02546009</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Brosnan</surname> <given-names>J. T.</given-names>
</name>
<name>
<surname>Brosnan</surname> <given-names>M. E.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>The sulfur-containing amino acids: an overview</article-title>. <source>J. Nutr.</source> <volume>136</volume>, <fpage>1636S</fpage>&#x2013;<lpage>1640S</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/jn/136.6.1636S</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="web">
<person-group person-group-type="author">
<collab>Canadian Grain Commission</collab>
</person-group> (<year>2022</year>) <source>Quality of Canadian oilseed-type soybeans</source>. Available at: <uri xlink:href="https://www.grainsCanada.gc.ca/en/grain-research/export-quality/oilseeds/soybean-oil/2018/pdf/report18.pdf">https://www.grainsCanada.gc.ca/en/grain-research/export-quality/oilseeds/soybean-oil/2018/pdf/report18.pdf</uri>.</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Zhong</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Ji</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Wan</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Shi</surname> <given-names>S.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Integrative analysis of metabolome and transcriptome reveals the improvements of seed quality in vegetable soybean (Glycine max (L.) Merr.)</article-title>. <source>Phytochemistry</source> <volume>200</volume>, <elocation-id>113216</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.phytochem.2022.113216</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chung</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Babka</surname> <given-names>H. L.</given-names>
</name>
<name>
<surname>Graef</surname> <given-names>G. L.</given-names>
</name>
<name>
<surname>Staswick</surname> <given-names>P. E.</given-names>
</name>
<name>
<surname>Lee</surname> <given-names>D. J.</given-names>
</name>
<name>
<surname>Cregan</surname> <given-names>P. B.</given-names>
</name>
<etal/>
</person-group>. (<year>2003</year>). <article-title>The seed protein, oil, and yield QTL on soybean linkage group I</article-title>. <source>Crop Sci.</source> <volume>43</volume>, <fpage>1053</fpage>&#x2013;<lpage>1067</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.2135/cropsci2003.1053</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Clemente</surname> <given-names>T. E.</given-names>
</name>
<name>
<surname>Cahoon</surname> <given-names>E. B.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Soybean oil: Genetic approaches for modification of functionality and total content</article-title>. <source>Plant Physiol.</source> <volume>151</volume>, <fpage>1030</fpage>&#x2013;<lpage>1040</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1104/pp.109.146282</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cober</surname> <given-names>E. R.</given-names>
</name>
<name>
<surname>Bing</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Voldeng</surname> <given-names>H. D.</given-names>
</name>
<name>
<surname>Soper</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Guillemette</surname> <given-names>R. J. D.</given-names>
</name>
<name>
<surname>Sloan</surname> <given-names>A.</given-names>
</name>
<etal/>
</person-group>. (<year>2006</year>). <article-title>90A01 soybean</article-title>. <source>Can. J. Plant Sci.</source> <volume>86</volume>, <fpage>481</fpage>&#x2013;<lpage>482</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4141/P05-187</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cober</surname> <given-names>E. R.</given-names>
</name>
<name>
<surname>Daba</surname> <given-names>K. A.</given-names>
</name>
<name>
<surname>Warkentin</surname> <given-names>T. D.</given-names>
</name>
<name>
<surname>Tomasiewicz</surname> <given-names>D. J.</given-names>
</name>
<name>
<surname>Mooleki</surname> <given-names>P. S.</given-names>
</name>
<name>
<surname>Karppinen</surname> <given-names>E. M.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Soybean seed protein content is lower but protein quality is higher in Western Canada compared with Eastern Canada</article-title>. <source>Can. J. Plant Sci</source> <volume>103</volume>(<issue>4</issue>), <fpage>411</fpage>&#x2013;<lpage>421</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1139/cjps-2022-0147</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Daley</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Deng</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Smith</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2020</year>) <source>The preseq Manual</source>. Available at: <uri xlink:href="http://smithlabresearch.org/manuals/preseqmanual.pdf">http://smithlabresearch.org/manuals/preseqmanual.pdf</uri>.</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>de Borja Reis</surname> <given-names>A. F.</given-names>
</name>
<name>
<surname>Tamagno</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Moro Rosso</surname> <given-names>L. H.</given-names>
</name>
<name>
<surname>Ortez</surname> <given-names>O. A.</given-names>
</name>
<name>
<surname>Naeve</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Ciampitti</surname> <given-names>I. A.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Historical trend on seed amino acid concentration does not follow protein changes in soybeans</article-title>. <source>Sci. Rep.</source> <volume>10</volume>, <fpage>1</fpage>&#x2013;<lpage>10</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41598-020-74734-1</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dembinski</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Bany</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>1991</year>). <article-title>The amino acid pool of high and low protein rye inbred lines (Secale cereale L.)</article-title>. <source>J. Plant Physiol.</source> <volume>138</volume>, <fpage>494</fpage>&#x2013;<lpage>496</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/S0176-1617(11)80529-8</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dobin</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Davis</surname> <given-names>C. A.</given-names>
</name>
<name>
<surname>Schlesinger</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Drenkow</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Zaleski</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Jha</surname> <given-names>S.</given-names>
</name>
<etal/>
</person-group>. (<year>2013</year>). <article-title>STAR: ultrafast universal RNA-seq aligner</article-title>. <source>Bioinformatics</source> <volume>29</volume>, <fpage>15</fpage>&#x2013;<lpage>21</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/bioinformatics/bts635</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fliege</surname> <given-names>C. E.</given-names>
</name>
<name>
<surname>Ward</surname> <given-names>R. A.</given-names>
</name>
<name>
<surname>Vogel</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Nguyen</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Quach</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Guo</surname> <given-names>M.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Fine mapping and cloning of the major seed protein quantitative trait loci on soybean chromosome 20</article-title>. <source>Plant J.</source> <volume>110</volume>, <fpage>114</fpage>&#x2013;<lpage>128</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/tpj.15658</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Galili</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>The aspartate-family pathway of plants: Linking production of essential amino acids with energy and stress regulation</article-title>. <source>Plant Signal. Behav.</source> <volume>6</volume>, <fpage>192</fpage>&#x2013;<lpage>195</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4161/psb.6.2.14425</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Garc&#xed;a-Alcalde</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Okonechnikov</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Carbonell</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Cruz</surname> <given-names>L. M.</given-names>
</name>
<name>
<surname>G&#xf6;tz</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Tarazona</surname> <given-names>S.</given-names>
</name>
<etal/>
</person-group>. (<year>2012</year>). <article-title>Qualimap: evaluating next-generation sequencing alignment data</article-title>. <source>Bioinformatics</source> <volume>28</volume>, <fpage>2678</fpage>&#x2013;<lpage>2679</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/bioinformatics/bts503</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gaufichon</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Rothstein</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Suzuki</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Asparagine metabolic pathways in arabidopsis</article-title>. <source>Plant Cell Physiol.</source> <volume>57</volume>, <elocation-id>pcv184</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/pcp/pcv184</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gomes</surname> <given-names>M. A. F.</given-names>
</name>
<name>
<surname>Sodek</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>1984</year>). <article-title>Allantoinase and asparaginase activities in maturing fruits of nodulated and non -nodulated soybeans</article-title>. <source>Physiol. Plant</source> <volume>62</volume>, <fpage>105</fpage>&#x2013;<lpage>109</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/j.1399-3054.1984.tb05931.x</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guti&#xe9;rrez</surname> <given-names>R. A.</given-names>
</name>
<name>
<surname>Stokes</surname> <given-names>T. L.</given-names>
</name>
<name>
<surname>Thum</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Obertello</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Katari</surname> <given-names>M. S.</given-names>
</name>
<etal/>
</person-group>. (<year>2008</year>). <article-title>Systems approach identifies an organic nitrogen-responsive gene network that is regulated by the master clock control gene CCA1</article-title>. <source>Proc. Natl. Acad. Sci. U. S. A.</source> <volume>105</volume>, <fpage>4939</fpage>&#x2013;<lpage>4944</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1073/pnas.0800211105</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Haga</surname> <given-names>K. I.</given-names>
</name>
<name>
<surname>Sodek</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>1987</year>). <article-title>Utilization of nitrogen sources by immature soybean cotyledons in culture</article-title>. <source>Ann. Bot.</source> <volume>59</volume>, <fpage>597</fpage>&#x2013;<lpage>601</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/oxfordjournals.aob.a087355</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Han</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Suglo</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Sun</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Su</surname> <given-names>T.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>L-aspartate: An essential metabolite for plant growth and stress acclimation</article-title>. <source>Molecules</source> <volume>26</volume>, <fpage>1</fpage>&#x2013;<lpage>17</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/molecules26071887</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Harada</surname> <given-names>J. J.</given-names>
</name>
<name>
<surname>Barker</surname> <given-names>S. J.</given-names>
</name>
<name>
<surname>Goldberg</surname> <given-names>R. B.</given-names>
</name>
</person-group> (<year>1989</year>). <article-title>Soybean beta-conglycinin genes are clustered in several DNA regions and are regulated by transcriptional and posttranscriptional processes</article-title>. <source>Plant Cell</source> <volume>1</volume>, <fpage>415</fpage>&#x2013;<lpage>425</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1105/tpc.1.4.415</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hern&#xe1;ndez-Sebasti&#xe0;</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Marsolais</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Saravitz</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Israel</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Dewey</surname> <given-names>R. E.</given-names>
</name>
<name>
<surname>Huber</surname> <given-names>S. C.</given-names>
</name>
</person-group> (<year>2005</year>). <article-title>Free amino acid profiles suggest a possible role for asparagine in the control of storage-product accumulation in developing seeds of low- and high-protein soybean lines</article-title>. <source>J. Exp. Bot.</source> <volume>56</volume>, <fpage>1951</fpage>&#x2013;<lpage>1963</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/jxb/eri191</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hooker</surname> <given-names>J. C.</given-names>
</name>
<name>
<surname>Nissan</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Luckert</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Charette</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Zapata</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Lefebvre</surname> <given-names>F.</given-names>
</name>
<etal/>
</person-group>. <article-title>A multi-year, multi-cultivar approach to differential expression analysis of high- and low-protein soybean (Glycine max)</article-title>. <source>Int. J. Mol. Sci.</source> (<year>2023</year>) <volume>24</volume>, <fpage>222</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/ijms24010222</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hooker</surname> <given-names>J. C.</given-names>
</name>
<name>
<surname>Nissan</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Luckert</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Zapata</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Hou</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Mohr</surname> <given-names>R. M.</given-names>
</name>
<etal/>
</person-group>. <article-title>GmSWEET29 and paralog GmSWEET34 are differentially expressed between soybeans grown in Eastern and Western Canada</article-title>. <source>Plants</source> (<year>2022</year>) <volume>11</volume>, <fpage>2337</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/plants11182337</pub-id>
</citation>
</ref> <ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Qi</surname> <given-names>H.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>Identification of soybean genes related to soybean seed protein content based on quantitative trait loci collinearity analysis</article-title>. <source>J. Agric. Food Chem.</source> <volume>67</volume>, <fpage>258</fpage>&#x2013;<lpage>274</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1021/acs.jafc.8b04602</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lam</surname> <given-names>H. M.</given-names>
</name>
<name>
<surname>Coschigano</surname> <given-names>K. T.</given-names>
</name>
<name>
<surname>Oliveira</surname> <given-names>I. C.</given-names>
</name>
<name>
<surname>Melo-Oliveira</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Coruzzi</surname> <given-names>G.M</given-names>
</name>
</person-group>(<year>1996</year>). and , <article-title>The molecular-genetics of nitrogen assimilation into amino acids in higher plants</article-title>. <source>Annu. Rev. Plant Physiol. Plant Mol. Biol.</source> <volume>47</volume>, <fpage>569</fpage>&#x2013;<lpage>593</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1146/annurev.arplant.47.1.569</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lam</surname> <given-names>H.-M.</given-names>
</name>
<name>
<surname>Wong</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Chan</surname> <given-names>H.-K.</given-names>
</name>
<name>
<surname>Yam</surname> <given-names>K.-M.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Chow</surname> <given-names>C.-M.</given-names>
</name>
<etal/>
</person-group>. (<year>2003</year>). <article-title>Overexpression of the ASN1 gene enhances nitrogen status in seeds of arabidopsis</article-title>. <source>Plant Physiol.</source> <volume>132</volume>, <fpage>926</fpage>&#x2013;<lpage>935</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1104/pp.103.020123</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Lea</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Miflin</surname> <given-names>B.</given-names>
</name>
</person-group> (<year>1980</year>). &#x201c;<article-title>Transport and metabolism of asparagine and other nitrogen compounds within the plant</article-title>,&#x201d; in <source>The biochemistry of plants</source>. Eds. <person-group person-group-type="editor">
<name>
<surname>Stumpt</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Conn</surname> <given-names>E.</given-names>
</name>
</person-group> (<publisher-loc>New York</publisher-loc>: <publisher-name>Academic Press</publisher-name>), <fpage>569</fpage>&#x2013;<lpage>607</lpage>.</citation>
</ref>
<ref id="B36">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Lea</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Robinson</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Stewart</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>1990</year>). &#x201c;<article-title>The enzymology and metabolism of glutamine, glutamate, and asparagine</article-title>,&#x201d; in <source>The biochemistry of plants: amino acids and derivatives</source>. Eds. <person-group person-group-type="editor">
<name>
<surname>Miflin</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Lea</surname> <given-names>P.</given-names>
</name>
</person-group> (<publisher-loc>New York</publisher-loc>: <publisher-name>Academic Press</publisher-name>), <fpage>121</fpage>&#x2013;<lpage>159</lpage>.</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lea</surname> <given-names>P. J.</given-names>
</name>
<name>
<surname>Sodek</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Parry</surname> <given-names>M. A. J.</given-names>
</name>
<name>
<surname>Shewry</surname> <given-names>P. R.</given-names>
</name>
<name>
<surname>Halford</surname> <given-names>N. G.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>Asparagine in plants</article-title>. <source>Ann. Appl. Biol.</source> <volume>150</volume>, <fpage>1</fpage>&#x2013;<lpage>26</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/j.1744-7348.2006.00104.x</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Liu</surname> <given-names>K.</given-names>
</name>
</person-group> (<year>1997</year>). <article-title>Chemistry and Nutritional Value of Soybean Components</article-title>. In. <source>Soybeans</source> (<publisher-loc>Boston, MA</publisher-loc>: <publisher-name>Springer</publisher-name>). doi:&#xa0;<pub-id pub-id-type="doi">10.1007/978-1-4615-1763-4_2</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lohaus</surname> <given-names>G.</given-names>
</name>
<name>
<surname>B&#xfc;ker</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Hu&#xdf;mann</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Soave</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Heldt</surname> <given-names>H.-W.</given-names>
</name>
</person-group> (<year>1998</year>). <article-title>Transport of amino acids with special emphasis on the synthesis and transport of asparagine in the Illinois Low Protein and Illinois High Protein strains of maize</article-title>. <source>Planta</source> <volume>205</volume>, <fpage>181</fpage>&#x2013;<lpage>188</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s004250050310</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Love</surname> <given-names>M. I.</given-names>
</name>
<name>
<surname>Huber</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Anders</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Moderated estimation of fold change and dispersion for RNA-seq data with DESeq2</article-title>. <source>Genome Biol.</source> <volume>15</volume>, <elocation-id>550</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s13059-014-0550-8</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Manhas</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Tripathi</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Khan</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Lakshmi</surname> <given-names>B. S.</given-names>
</name>
<name>
<surname>Lal</surname> <given-names>S. K.</given-names>
</name>
<name>
<surname>Gowri</surname> <given-names>V. S.</given-names>
</name>
<etal/>
</person-group>. (<year>2014</year>). <article-title>Identification and functional characterization of a novel bacterial type asparagine synthetase A: A trna synthetase paralog from leishmania donovani</article-title>. <source>J. Biol. Chem.</source> <volume>289</volume>, <fpage>12096</fpage>&#x2013;<lpage>12108</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1074/jbc.M114.554642</pub-id>
</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Morley</surname> <given-names>S. A.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Alazem</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Frankfater</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Yi</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Burch-Smith</surname> <given-names>T.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Expression of Malic enzyme reveals subcellular carbon partitioning for storage reserve production in soybeans</article-title>. <source>New Phytol</source> <volume>239</volume>, <fpage>1834</fpage>&#x2013;<lpage>1851</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/nph.18835</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Natarajan</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Luthria</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Bae</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Lakshman</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Mitra</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>Transgenic soybeans and soybean protein analysis: an overview</article-title>. <source>J. Agric. Food Chem.</source> <volume>61</volume>, <fpage>11736</fpage>&#x2013;<lpage>11743</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1021/jf402148e</pub-id>
</citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ort</surname> <given-names>N. W. W.</given-names>
</name>
<name>
<surname>Morrison</surname> <given-names>M. J.</given-names>
</name>
<name>
<surname>Cober</surname> <given-names>E. R.</given-names>
</name>
<name>
<surname>McAndrew</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Lawley</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>A comparison of soybean maturity groups for phenology, seed yield, and seed quality components between eastern Ontario and southern Manitoba</article-title>. <source>Can. J. Plant Sci</source> <volume>102</volume>(<issue>4</issue>), <fpage>812</fpage>&#x2013;<lpage>822</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1139/CJPS-2021-0235</pub-id>
</citation>
</ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pandurangan</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Pajak</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Molnar</surname> <given-names>S. J.</given-names>
</name>
<name>
<surname>Cober</surname> <given-names>E. R.</given-names>
</name>
<name>
<surname>Dhaubhadel</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Hern&#xe1;ndez-Sebasti</surname> <given-names>C.</given-names>
</name>
<etal/>
</person-group>. (<year>2012</year>). <article-title>Relationship between asparagine metabolism and protein concentration in soybean seed</article-title>. <source>J. Exp. Bot.</source> <volume>63</volume>, <fpage>3173</fpage>&#x2013;<lpage>3184</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/jxb/ers039</pub-id>
</citation>
</ref>
<ref id="B46">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Pedersen</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Licht</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2014</year>). <source>Soybean growth and development</source> (<publisher-loc>Ames, USA</publisher-loc>: <publisher-name>Iowa State University Extension</publisher-name>). PM 1945.</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Peng</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Qian</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Song</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Cheng</surname> <given-names>H.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>Comparative transcriptome analysis during seeds development between two soybean cultivars</article-title>. <source>PeerJ</source> <volume>9</volume>, <fpage>1</fpage>&#x2013;<lpage>20</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.7717/peerj.10772</pub-id>
</citation>
</ref>
<ref id="B48">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Peng</surname> <given-names>I. C.</given-names>
</name>
<name>
<surname>Quass</surname> <given-names>D. W.</given-names>
</name>
<name>
<surname>Dayton</surname> <given-names>W. R.</given-names>
</name>
<name>
<surname>Allen</surname> <given-names>C. E.</given-names>
</name>
</person-group> (<year>1984</year>). <article-title>Physico chemical properties of soybean 11S globulin-A Review</article-title>. <source>Cereal Chem.</source> <volume>61</volume>, <fpage>480</fpage>&#x2013;<lpage>490</lpage>.</citation>
</ref>
<ref id="B49">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rainbird</surname> <given-names>R. M.</given-names>
</name>
<name>
<surname>Thorne</surname> <given-names>J. H.</given-names>
</name>
<name>
<surname>Hardy</surname> <given-names>R. W. F.</given-names>
</name>
</person-group> (<year>1984</year>). <article-title>Role of amides, amino acids, and ureides in the nutrition of developing soybean seeds</article-title>. <source>Plant Physiol.</source> <volume>74</volume>, <fpage>329</fpage>&#x2013;<lpage>334</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1104/pp.74.2.329</pub-id>
</citation>
</ref>
<ref id="B50">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Robinson</surname> <given-names>M. D.</given-names>
</name>
<name>
<surname>McCarthy</surname> <given-names>D. J.</given-names>
</name>
<name>
<surname>Smyth</surname> <given-names>G. K.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>edgeR: a Bioconductor package for differential expression analysis of digital gene expression data</article-title>. <source>Bioinformatics</source> <volume>26</volume>, <fpage>139</fpage>&#x2013;<lpage>140</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/bioinformatics/btp616</pub-id>
</citation>
</ref>
<ref id="B51">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sayols</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Scherzinger</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Klein</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>dupRadar: A Bioconductor package for the assessment of PCR artifacts in RNA-Seq data</article-title>. <source>BMC Bioinf.</source> <volume>17</volume>, <fpage>1</fpage>&#x2013;<lpage>5</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s12859-016-1276-2</pub-id>
</citation>
</ref>
<ref id="B52">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Servet</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Ghelis</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Richard</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Zilberstein</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Savoure</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Proline dehydrogenase: a key enzyme in controlling cellular homeostasis</article-title>. <source>FBL</source> <volume>17</volume>, <fpage>607</fpage>&#x2013;<lpage>620</lpage>. doi: <pub-id pub-id-type="doi">10.2741/3947</pub-id>
</citation>
</ref>
<ref id="B53">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Severin</surname> <given-names>A. J.</given-names>
</name>
<name>
<surname>Woody</surname> <given-names>J. L.</given-names>
</name>
<name>
<surname>Bolon</surname> <given-names>Y.-T.</given-names>
</name>
<name>
<surname>Joseph</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Diers</surname> <given-names>B. W.</given-names>
</name>
<name>
<surname>Farmer</surname> <given-names>A. D.</given-names>
</name>
<etal/>
</person-group>. (<year>2010</year>). <article-title>RNA-Seq Atlas of Glycine max: A guide to the soybean transcriptome</article-title>. <source>BMC Plant Biol.</source> <volume>10</volume>, <elocation-id>160</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/1471-2229-10-160</pub-id>
</citation>
</ref>
<ref id="B54">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Stephens</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>False discovery rates: a new deal</article-title>. <source>Biostatistics</source> <volume>18</volume>, <fpage>275</fpage>&#x2013;<lpage>294</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/biostatistics/kxw041</pub-id>
</citation>
</ref>
<ref id="B55">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tsukada</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Kitamura</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Harada</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Kaizuma</surname> <given-names>N.</given-names>
</name>
</person-group> (<year>1986</year>). <article-title>Genetic analysis of subunits of two major storage proteins (&#x3b2;-conglycinin and glycinin) in soybean seeds</article-title>. <source>Japanese J. Breed.</source> <volume>36</volume>, <fpage>390</fpage>&#x2013;<lpage>400</lpage>. doi: <pub-id pub-id-type="doi">10.1270/jsbbs1951.36.390</pub-id>
</citation>
</ref>
<ref id="B56">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vinogradova</surname> <given-names>I. S.</given-names>
</name>
<name>
<surname>Falaleev</surname> <given-names>O. V.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Formation of the vascular system of developing bean (Phaseolus limensis L.) seeds according to nuclear magnetic resonance microtomography</article-title>. <source>Russ. J. Dev. Biol.</source> <volume>43</volume>, <fpage>25</fpage>&#x2013;<lpage>34</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1134/S1062360412010079</pub-id>
</citation>
</ref>
<ref id="B57">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Voldeng</surname> <given-names>H. D.</given-names>
</name>
<name>
<surname>Guillemette</surname> <given-names>R. J. D.</given-names>
</name>
<name>
<surname>Leonard</surname> <given-names>D. A.</given-names>
</name>
<name>
<surname>Cober</surname> <given-names>E. R.</given-names>
</name>
</person-group> (<year>1996</year>a). <article-title>AC harmony soybean</article-title>. <source>Can. J. Plant Sci.</source> <volume>76</volume>, <fpage>477</fpage>&#x2013;<lpage>478</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4141/cjps96-086</pub-id>
</citation>
</ref>
<ref id="B58">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Voldeng</surname> <given-names>H. D.</given-names>
</name>
<name>
<surname>Guillemette</surname> <given-names>R. J. D.</given-names>
</name>
<name>
<surname>Leonard</surname> <given-names>D. A.</given-names>
</name>
<name>
<surname>Cober</surname> <given-names>E. R.</given-names>
</name>
</person-group> (<year>1996</year>b). <article-title>AC proteus soybean</article-title>. <source>Can. J. Plant Sci.</source> <volume>76</volume>, <fpage>153</fpage>&#x2013;<lpage>154</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4141/cjps96-031</pub-id>
</citation>
</ref>
<ref id="B59">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Voorrips</surname> <given-names>R. E.</given-names>
</name>
</person-group> (<year>2002</year>). <article-title>MapChart: software for the graphical presentation of linkage maps and QTLs</article-title>. <source>J. Hered.</source> <volume>93</volume>, <fpage>77</fpage>&#x2013;<lpage>78</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/jhered/93.1.77</pub-id>
</citation>
</ref>
<ref id="B60">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wan</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Shao</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Shan</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Zeng</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Lam</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>Correlation between AS1 gene expression and seed protein contents in different soybean (Glycine max [L.] merr.) cultivars</article-title>. <source>Plant Biol.</source> <volume>8</volume>, <fpage>271</fpage>&#x2013;<lpage>276</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1055/s-2006-923876</pub-id>
</citation>
</ref>
<ref id="B61">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Yokosho</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>Y.-C.</given-names>
</name>
<etal/>
</person-group>. (<year>2020</year>). <article-title>Simultaneous changes in seed size, oil content and protein content driven by selection of SWEET homologues during soybean domestication</article-title>. <source>Natl. Sci. Rev.</source> <volume>7</volume>, <fpage>1776</fpage>&#x2013;<lpage>1786</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/nsr/nwaa110</pub-id>
</citation>
</ref>
<ref id="B62">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>W.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>RSeQC: quality control of RNA-seq experiments</article-title>. <source>Bioinformatics</source> <volume>28</volume>, <fpage>2184</fpage>&#x2013;<lpage>2185</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/bioinformatics/bts356</pub-id>
</citation>
</ref>
<ref id="B63">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Shi</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Yan</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>Q.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>Primary metabolite contents are correlated with seed protein and oil traits in near-isogenic lines of soybean</article-title>. <source>Crop J.</source> <volume>7</volume>, <fpage>651</fpage>&#x2013;<lpage>659</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.cj.2019.04.002</pub-id>
</citation>
</ref>
<ref id="B64">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Whitfield</surname> <given-names>C. D.</given-names>
</name>
<name>
<surname>Steers</surname> <given-names>E. J.</given-names>
</name>
<name>
<surname>Weisbach</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>1970</year>). <article-title>Purification and properties of 5-methyltetrahydropteroyltriglutamate-homocysteine transmethylase</article-title>. <source>J. Biol. Chem.</source> <volume>245</volume>, <fpage>390</fpage>&#x2013;<lpage>401</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/s0021-9258(18)63404-0</pub-id>
</citation>
</ref>
<ref id="B65">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Song</surname> <given-names>Q.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Using transcriptomic and metabolomic data to investigate the molecular mechanisms that determine protein and oil contents during seed development in soybean</article-title>. <source>Front. Plant Sci.</source> <volume>13</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fpls.2022.1012394</pub-id>
</citation>
</ref>
<ref id="B66">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Lee</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Pandurangan</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Clarke</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Pajak</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Marsolais</surname> <given-names>F.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>Characterization of Arabidopsis serine:glyoxylate aminotransferase, AGT1, as an asparagine aminotransferase</article-title>. <source>Phytochemistry</source> <volume>85</volume>, <fpage>30</fpage>&#x2013;<lpage>35</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.phytochem.2012.09.017</pub-id>
</citation>
</ref>
<ref id="B67">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Lu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Bhusal</surname> <given-names>S. J.</given-names>
</name>
<name>
<surname>Song</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Cregan</surname> <given-names>P. B.</given-names>
</name>
<etal/>
</person-group>. (<year>2018</year>). <article-title>Genome-wide scan for seed composition provides insights into soybean quality improvement and the impacts of domestication and breeding</article-title>. <source>Mol. Plant</source> <volume>11</volume>, <fpage>460</fpage>&#x2013;<lpage>472</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.molp.2017.12.016</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>