<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Plant Sci.</journal-id>
<journal-title>Frontiers in Plant Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Plant Sci.</abbrev-journal-title>
<issn pub-type="epub">1664-462X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpls.2025.1666460</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Plant Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Near-infrared prediction of total phosphorus in leaves content in korla fragrant pear with growth period specificity via spectral modeling</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Yu</surname>
<given-names>Mingyang</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="author-notes" rid="fn003">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2987837/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Zeng</surname>
<given-names>Junkai</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="author-notes" rid="fn003">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/3134140/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Li</surname>
<given-names>Yang</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Fan</surname>
<given-names>Weifan</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Wang</surname>
<given-names>Lanfei</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Wang</surname>
<given-names>Hao</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Bao</surname>
<given-names>Jianping</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>*</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Xinjiang Production and Construction Corps, Tarim University Institute of Horticulture and Forestry</institution>, <addr-line>Xinjiang Alar</addr-line>,&#xa0;<country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Tarim Basin Biological Resources Protection and Utilization Key Laboratory, Xinjiang Production and Construction Corps</institution>, <addr-line>Xinjiang Alar</addr-line>,&#xa0;<country>China</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Xinjiang Production and Construction Corps, Southern Xinjiang Special Fruit Trees High -Quality, High -Quality Cultivation and Deep Processing of Fruit Products Processing Technical National Local Joint Engineering Laboratory</institution>, <addr-line>Xinjiang Alar</addr-line>,&#xa0;<country>China</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Ministry of Education of the People&#x2019;s Republic of China, Nanjing Agricultural University Horticulture and Forestry College</institution>, <addr-line>Nanjing, Jiangsu</addr-line>,&#xa0;<country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>Edited by: <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1089966/overview">Bijayalaxmi Mohanty</ext-link>, National University of Singapore, Singapore</p>
</fn>
<fn fn-type="edited-by">
<p>Reviewed by: <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1907515/overview">Shenghui Yang</ext-link>, China Agricultural University, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/3143474/overview">Xin Wu</ext-link>, Chongqing Polytechnic University of Electronic Technology, China</p>
</fn>
<fn fn-type="corresp" id="fn001">
<p>*Correspondence: Jianping Bao, <email xlink:href="mailto:baobao-xinjiang@126.com">baobao-xinjiang@126.com</email>
</p>
</fn>
<fn fn-type="equal" id="fn003">
<p>&#x2020;These authors have contributed equally to this work</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>10</day>
<month>10</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>16</volume>
<elocation-id>1666460</elocation-id>
<history>
<date date-type="received">
<day>15</day>
<month>07</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>29</day>
<month>09</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2025 Yu, Zeng, Li, Fan, Wang, Wang and Bao.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Yu, Zeng, Li, Fan, Wang, Wang and Bao</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Leaf total phosphorus content (LTP) is a key indicator for assessing fruit nutrition status. As a rapid non-destructive inspection method, Near-infrared spectroscopy technology is susceptible to the influence of changes in plant growth periods and spectral noise on its prediction accuracy. At present, how to synergistically utilize growth period information and Spectral pre - processing methods to optimize the LTP Prediction model remains to be further studied. The study systematically collected Leaf sample and their near-infrared Spectral data during three key growth periods of Korla fragrant pear (fruit-setting period, fruit swelling period, and Maturity period). In the Spectral pre-processing stage, multiple scattering correction, Savitzky-Golay Smooth, First Derivative (FD), Second Derivative (SD) and their combined algorithms were comprehensively applied. The Competitive Adaptive Reweighted Sampling (CARS) algorithm was used for characteristic wavelength selection, and based on this, Growth period specificity BP neural network model and cross-growth period general prediction models were constructed respectively to evaluate the performance of different Modeling strategies. Results The study showed that LTP content exhibited a significant differential distribution across different growing stage. In the characteristic wavelength bands, after processing with Combined pre-processing method (e.g., MSC+ FD), the correlation coefficient between the spectrum and LTP content significantly increased to approximately 0.90. The predictive performance of the Growth-period-specific model was comprehensively superior to that of the general model, with the Validation set coefficient of determination remaining above 0.83. Compared with the general model, the Coefficient of determination (R<sup>2</sup>) increased by 0.05-0.16, and the root mean square error decreased by 0.0029-0.0079. This study successfully constructed a technical system of &#x201c;Growth period-Preprocessing-Model&#x201d;. The results indicated that the Modeling strategy considering the characteristics of crop growing stage could significantly improve the predictive ability of near-infrared spectroscopy models. This study provides a reliable technical framework for Precision nutrient management in orchard, and the established methodology can also serve as a reference for nutrient Surveillance of other fruit tree plants.</p>
</abstract>
<kwd-group>
<kwd>korla fragrant pear</kwd>
<kwd>total phosphorus in leaves</kwd>
<kwd>near-infrared spectrum</kwd>
<kwd>growth period specificity</kwd>
<kwd>machine learning</kwd>
</kwd-group>
<counts>
<fig-count count="10"/>
<table-count count="4"/>
<equation-count count="4"/>
<ref-count count="64"/>
<page-count count="20"/>
<word-count count="10094"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-in-acceptance</meta-name>
<meta-value>Plant Biophysics and Modeling</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1" sec-type="intro">
<label>1</label>
<title>Introduction</title>
<p>Phosphorus as an essential mineral element for plant growth and development plays a critical role in physiological processes such as nucleic acid synthesis, energy metabolism, and the maintenance of cell membrane structures (<xref ref-type="bibr" rid="B9">Chen et&#xa0;al., 2021</xref>; <xref ref-type="bibr" rid="B41">Song et&#xa0;al., 2024</xref>).The dynamic change of Leaf total phosphorus (LTP) content is not only a direct reflection of fruit nutrition status, but also an important basis for precise fertilization in orchards (<xref ref-type="bibr" rid="B38">Shah et&#xa0;al., 2024</xref>).As a characteristic cash crop in the arid regions of northwestern China, Phosphorus nutritional diagnosis in korla fragrant pear has significant practical implications for improving fruit quality and yield (<xref ref-type="bibr" rid="B48">Wang et&#xa0;al., 2022</xref>).</p>
<p>Near-infrared spectroscopy (NIRS) technology offers an innovative approach for <italic>in-situ</italic> monitoring of nutritive element of plant due to its advantages of non-destructive inspection, high-throughput analysis, and rapid response. By capturing the vibrational absorption features of hydrogen-containing groups (e.g., P-O-H), it enables spectrum analysis of leaf Phosphorus content (<xref ref-type="bibr" rid="B33">Murguzur et&#xa0;al., 2019</xref>). </p>
<p>Current research on fruit tree Phosphorus Spectral diagnosis faces three bottlenecks that require breakthroughs: First, most studies have not systematically considered the impact of Growth period differences on leaf Phosphorus distribution. The phosphorus metabolism characteristics of korla fragrant pear differ significantly during the fruit-setting period, fruit-expanding period, and maturity, exhibiting distribution patterns of low content and high dispersion during the fruit-setting period, stable state during the fruit-expanding period, and high content and high dispersion during maturity (<xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>). These patterns necessitate model construction that adapts to the physiological characteristics of different phenological periods. However, existing studies mostly adopt an intertemporal general model, limiting prediction accuracy (<xref ref-type="bibr" rid="B40">Siedliska et&#xa0;al., 2021</xref>).Second, the synergistic mechanism of Spectral preprocessing technology remains unclear. Although single preprocessing methods (such as Multivariate scattering correction MSC, Derivative processing FD/SD) can separately achieve physical interference elimination or chemical characteristics enhancement, they struggle to simultaneously meet the dual requirements of noise suppression and dynamic information preservation. Systematic exploration of optimized combined preprocessing strategies is still lacking (<xref ref-type="bibr" rid="B22">Han et&#xa0;al., 2025</xref>; <xref ref-type="bibr" rid="B34">Qi et&#xa0;al., 2025</xref>).Third, Feature band selection and Model parameter optimization are not dynamically coupled with the Growth period. The spectral absorption peak associated with Phosphorus (4000&#x2013;7500 cm<sup>-</sup>&#xb9;) exhibits significant differences in response intensity across different Growth period, whereas traditional feature selection algorithms fail to fully exploit this time-space specificity, resulting in insufficient model generalization ability (<xref ref-type="bibr" rid="B45">Tian et&#xa0;al., 2024</xref>).</p>
<fig id="f1" position="float">
<label>Figure&#xa0;1</label>
<caption>
<p>Experimental overall visualization flowchart.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g001.tif">
<alt-text content-type="machine-generated">Flowchart illustrating the process of test sample handling and image visualization. It includes sections on sample collection, raw spectral acquisition, potassium determination, and model evaluation methods. The image visualization part comprises raw spectrum graphs, spectral preprocessing and correlation, spectral feature extraction, model parameter visualization, and various charts and graphs depicting data analysis and modeling results.</alt-text>
</graphic>
</fig>
<fig id="f2" position="float">
<label>Figure&#xa0;2</label>
<caption>
<p>Content of korla fragrant pear total phosphorus in leaves in different periods; Different letters indicate significant differences between groups(P&lt; 0.05).The black dots represent the LTP content of each sample within each Growth period, and the black curve represents the foot normal distribution.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g002.tif">
<alt-text content-type="machine-generated">Box plots showing LTP percentage during three phases: fruit setting (lowest percentage), fruit enlargement (mid percentage), and fruit ripening (highest percentage). Significant differences are marked by letters a, b, c.</alt-text>
</graphic>
</fig>
<p>It is worth noting that although near-infrared spectrum analysis has been widely applied in the non-destructive detection of crop nutrients, research on quantitative prediction of phosphorus in fruit trees remains insufficient. Most current studies have focused on field crops such as wheat (<xref ref-type="bibr" rid="B60">Zhang et&#xa0;al., 2022</xref>) and rice (<xref ref-type="bibr" rid="B4">Arias et&#xa0;al., 2021</xref>), with inadequate exploration of the relationship between leaf total phosphorus (LTP) content and spectral response in fruit tree leaves such as korla fragrant pear. Due to the relatively complex morphological structure of fruit tree leaves, combined with variations in the canopy microenvironment and physiological dynamics at different growing stages, the difficulty of Spectral modeling is increased. Existing research methods often directly apply Traditional regression algorithms such as PLSR and SVR (<xref ref-type="bibr" rid="B1">Ahmadi et&#xa0;al., 2021</xref>), failing to conduct targeted model improvement based on the spectral characteristics of fruit trees, and particularly lacking a systematic research approach that integrates phenological change, preprocessing method, and machine learning model. Therefore, establishing a spectral prediction model for LTP content that can respond to the Growth period characteristics of fragrant pear is of great significance for achieving precise monitoring of phosphorus nutrition.</p>
<p>To address the above research bottlenecks, this study aims to overcome the limitations of traditional general models and achieve systematic innovation from theoretical, technical, and applied perspectives, specifically reflected in: 1. Systematically analyzing the unique distribution patterns (left-skewed, stable, right-skewed) of LTP content in korla fragrant pear at different Growth periods (<xref ref-type="bibr" rid="B13">Fonseca-Garc&#xed;a et&#xa0;al., 2021</xref>)and their differential requirements for spectral models, providing a solid physiological basis for Stage-based modeling.2. In-depth exploration of various preprocessing methods (single and combined) under different growth periods within the Collaborative optimization mechanism (e.g., MSC+FD for high-dispersion stages, SG+SD for weak-signal stages), rather than simple stacking, to achieve efficient spectral information purification (<xref ref-type="bibr" rid="B44">Tan et&#xa0;al., 2025</xref>). 3. Construction of a complete technical system of &#x201c;growth period specificity-preprocessing collaboration-model adaptation&#x201d; to validate the performance improvement of the Stage-based modeling strategy compared to a general model (<xref ref-type="bibr" rid="B7">Cao et&#xa0;al., 2021</xref>), providing a directly applicable solution for precision orchard management.</p>
<p>To this end, this study first analyzes the distribution characteristics of leaf total phosphorus (LTP) in korla fragrant pear across different growth periods using Box plot, clarifying the content dynamics during the fruit-setting period (minimum 0.02%, maximum 0.25%, left-skewed distribution), fruit-expanding period (median 0.15%, concentrated in 0.10%&#x2013;0.20%), and maturity period (maximum 0.45%, right-skewed distribution) (<xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>), providing a physiological basis for Spectral modeling; secondly, integrating MSC, SG smoothing, FD/SD derivative processing, and combined strategies (MSC+FD, SG+SD, etc.) to optimize spectral signals, achieving synergy between physical interference elimination and chemical characteristics enhancement in the core sensitive region of 4000&#x2013;5500 cm<sup>-</sup>&#xb9; and the 5500&#x2013;7500 cm<sup>-</sup>&#xb9; combined frequency region (<xref ref-type="fig" rid="f3">
<bold>Figures&#xa0;3</bold>
</xref>-<xref ref-type="fig" rid="f4">
<bold>4</bold>
</xref>); further, using the competitive adaptive reweighted sampling (CARS) algorithm to screening feature band (<xref ref-type="bibr" rid="B58">Zhang et&#xa0;al., 2023</xref>)2. In-depth exploration of various preprocessing methods (single and combined) under different growth periods within the Collaborative optimization mechanism (e.g., MSC+FD for high-dispersion stages, SG+SD for weak-signal stages), rather than simple stacking, to achieve efficient spectral information purification (<xref ref-type="bibr" rid="B44">Tan et&#xa0;al., 2025</xref>). 3. Construction of a complete technical system of &#x201c;growth period specificity-preprocessing collaboration-model adaptation&#x201d; to validate the performance improvement of the Stage-based modeling strategy compared to a general model (<xref ref-type="bibr" rid="B7">Cao et&#xa0;al., 2021</xref>), providing a directly applicable solution for precision orchard management.</p>
<fig id="f3" position="float">
<label>Figure&#xa0;3</label>
<caption>
<p>Analysis of original spectral images: <bold>(A)</bold> is the original spectral image; <bold>(B)</bold> is the interior visualization of the average of human spectral images over different periods.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g003.tif">
<alt-text content-type="machine-generated">Graph A shows multiple spectral lines depicting absorbance versus wavenumber, ranging from 4000 to 10000 inverse centimeters. Graph B illustrates clustered spectral lines representing fruit setting, enlargement, and ripening periods, distinguishing between them by color. Both graphs measure absorbance from 0.2 to 1.6.</alt-text>
</graphic>
</fig>
<fig id="f4" position="float">
<label>Figure&#xa0;4</label>
<caption>
<p>Different spectral images under different preprocessing methods: <bold>(A)</bold>: MSC; <bold>(B)</bold>: SG; <bold>(C)</bold>: FD; <bold>(D)</bold>: SD; <bold>(E)</bold>: MSC+FD; <bold>(F)</bold>: MSC+SD; <bold>(G)</bold>: SG+FD; <bold>(H)</bold>: SG+SD.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g004.tif">
<alt-text content-type="machine-generated">Eight graphs labeled A to H illustrate absorbance versus wavenumber from 4000 to 10000 cm&#x207b;&#xb9;. Graphs A and B show higher absorbance values, while C and D display lower fluctuations. Graphs E to H depict minimized absorbance variations, with a focus on detailed wave patterns. Each graph varies in line density and absorbance scale.</alt-text>
</graphic>
</fig>
<p>To this end, this study first analyzes the distribution characteristics of leaf total phosphorus (LTP) in korla fragrant pear across different growth periods using Box plot, clarifying the content dynamics during the fruit-setting period (minimum 0.02%, maximum 0.25%, left-skewed distribution), fruit-expanding period (median 0.15%, concentrated in 0.10%&#x2013;0.20%), and maturity period (maximum 0.45%, right-skewed distribution) (<xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>), providing a physiological basis for Spectral modeling; secondly, integrating MSC, SG smoothing, FD/SD derivative processing, and combined strategies (MSC+FD, SG+SD, etc.) to optimize spectral signals, achieving synergy between physical interference elimination and chemical characteristics enhancement in the core sensitive region of 4000&#x2013;5500 cm<sup>-</sup>&#xb9; and the 5500&#x2013;7500 cm<sup>-</sup>&#xb9; combined frequency region (<xref ref-type="fig" rid="f3">
<bold>Figures&#xa0;3</bold>
</xref>-<xref ref-type="fig" rid="f4">
<bold>4</bold>
</xref>); further, using the competitive adaptive reweighted sampling (CARS) algorithm to screening feature band (<xref ref-type="bibr" rid="B58">Zhang et&#xa0;al., 2023</xref>), combined with algorithms such as BP neural network (<xref ref-type="bibr" rid="B55">Yang et&#xa0;al., 2021</xref>)and random forest (<xref ref-type="bibr" rid="B8">Capitaine et&#xa0;al., 2021</xref>)to construct Growth-period-specific model, with model performance evaluated and compared through metrics including Coefficient of determination (R&#xb2;) and Root mean square error (RMSE).</p>
<p>This study not only improves the Growth period adaptation theory for Spectral diagnosis of fruit tree Phosphorus, but also provides a Methodological reference for the application of Near-infrared technology in Precise orchard nutrient management. As an important component of a series of studies, this result corroborates previous Spectral diagnosis research on plants such as potato and rice (<xref ref-type="bibr" rid="B59">Zhang et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B14">Gao et&#xa0;al., 2023</xref>), collectively revealing the coupling pattern of &#x2018;spectral trait-growth period-nutritive index&#x2019;, thus laying the foundation for constructing a universal technical system for fruit nutrition diagnosis (<xref ref-type="bibr" rid="B50">Xiao et&#xa0;al., 2022</xref>; <xref ref-type="bibr" rid="B17">Guo et&#xa0;al., 2024</xref>). Subsequent research will focus on the integration of Mid-infrared spectroscopy and near-infrared spectrum, as well as correction mechanisms for field environmental interferences, promoting the translation of Spectral diagnosis technology from laboratory research to practical application.</p>
</sec>
<sec id="s2" sec-type="materials|methods">
<label>2</label>
<title>Materials and methods</title>
<sec id="s2_1">
<label>2.1</label>
<title>Overview of the test site</title>
<p>The experiment was conducted at the campus of Tarim University in Alar City, Xinjiang. The test material was 23-year-old Korla Fragrant Pear (grafted onto Birchleaf Pear rootstock), planted in north-south rows with a spacing of 2 m&#xd7;4 m. The orchard was irrigated using the Flood irrigation method, and other management practices were carried out according to local conventional protocols. Mature trees with vigorous growth and uniform tree vigor were selected for the study. Please refer to <xref ref-type="fig" rid="f1"><bold>Figure&#xa0;1</bold></xref> for the experiment process.</p>
</sec>
<sec id="s2_2">
<label>2.2</label>
<title>Sample collection</title>
<p>At the fruit bearing periods (April 23, 2024), fruit expansion period (July 11, 2024), and maturity period (September 20, 2024) of Kuerle fragrant pear fruit, mature leaves were collected from the middle and lower segments of current-year branches at the outer edge of the tree crown of each Test tree. During collection, single leaves from the east, south, west, and north directions of the tree crown were carefully selected. Leaves from 150 trees were collected for each period, labeled, and stored in Ziplock bag inside a 4&#xb0;C refrigerator for subsequent Spectral scanning and Total phosphorus content analysis.</p>
</sec>
<sec id="s2_3">
<label>2.3</label>
<title>Original spectrum acquisition</title>
<p>Remove Test sample from the -4&#xb0;C freezer and place it in the laboratory where the spectrometer is located (ambient temperature 24&#xb0;C) for 12 hours to equilibrate, ensuring that the sample temperature is consistent with room temperature and eliminating interference from temperature gradients (<xref ref-type="bibr" rid="B61">Zheng et&#xa0;al., 2023</xref>). After powering on the Fourier Transform Near-Infrared Spectrometer (Antaris II FT-NIR) and allowing it to warm up for 30 minutes, perform Diffuse reflectance correction using the Standard whiteboard (<xref ref-type="bibr" rid="B23">He et&#xa0;al., 2023</xref>). If necessary, gently clean dust from the leaf surface with a dust-free cloth. On the leaf, select two regions each at the upper and lower ends, using the vein as a boundary (four sites in total), and use different colors to mark the spectra of different regions. At each region, repeat the scan 4 times using the following parameters: Spectral range 10000&#x2013;4000 cm<sup>-</sup>&#xb9;, Resolution 8 cm<sup>-</sup>&#xb9;, Gain 2&#xd7;, Number of accumulations per scan 64 times (<xref ref-type="bibr" rid="B62">Zheng et&#xa0;al., 2024</xref>). Single leaf yielded 16 spectral curves, and after baseline correction, the average was calculated and used as the final Absorbance (A) value of the sample for subsequent Chemometric modeling and analysis. This method effectively controlled the effects of Temperature fluctuation, Instrument drift, and Leaf heterogeneity through Standardized preprocessing, Instrument calibration, and Multi-point repeated measurement, laying the data foundation for constructing a High-precision prediction model.</p>
</sec>
<sec id="s2_4">
<label>2.4</label>
<title>Determination of total potassium in korla fragrant pear leaves</title>
<p>Collect the Leaf sample after Spectral data acquisition and sequentially wash with tap water, 0.1% detergent solution, tap water, and Distilled water (entire process &#x2264; 2 min). After removing surface moisture with Dust-free absorbent paper, place the sample in a 105&#xb0;C Forced air drying oven for fixation for 20 min, then dry to constant weight at 80&#xb0;C (<xref ref-type="bibr" rid="B12">Dayton et&#xa0;al., 2017</xref>); grind the dried sample using a Stainless steel crusher and pass through a 60-mesh nylon sieve (<xref ref-type="bibr" rid="B53">Xu et&#xa0;al., 2016</xref>). Accurately weigh 0.2000 g of the Sieved sample into a 100 mL Digestion tube. Moisten the sample with Distilled water, then add 5 mL of Concentrated sulfuric acid (H<sub>2</sub>SO<sub>4</sub>), and fit the mouth of the tube with a curved neck funnel. On a Digestion furnace, initially heat gently at low temperature; gradually increase the temperature after dense white smoke appears due to decomposition of sulfuric acid. When the solution turns brown-black, remove the Digestion tube, cool, then add 10 drops of 300 g&#xb7;L<sup>-</sup>&#xb9; hydrogen peroxide (H<sub>2</sub>O<sub>2</sub>) dropwise while thoroughly shaking. Continue heating for 15 min. Repeat the above H<sub>2</sub>O<sub>2</sub> addition operation 2&#x2013;3 times until the Digestion solution becomes colorless or clear and transparent, then heat for an additional 10 min to completely remove excess H<sub>2</sub>O<sub>2</sub> (<xref ref-type="bibr" rid="B64">Zou et&#xa0;al., 2019</xref>); after cooling, rinse the curved neck funnel with Distilled water, combine the rinsing solution into the Digestion tube, and dilute to the 100 mL mark; determination of Total phosphorus content is performed using the molybdenum antimony resistance colorimetric method (<xref ref-type="bibr" rid="B30">Li Ting and Hong Xing, 2022</xref>), specifically: transfer 5 mL of Digestion solution (pre-dilute if concentration is too high) into a 50 mL volumetric flask, add 5 mL of freshly prepared molybdenum-antimony-ascorbic acid color developer [preparation: slowly add 100 mL of 0.5 mol/L H<sub>2</sub>SO<sub>4</sub> to a mixed solution containing 10 g ammonium molybdate and 0.5 g antimony potassium tartrate, cool, then add 1.5 g of Ascorbic acid and dilute to 500 mL], dilute to volume with Distilled water, and allow color development in the dark at 20&#x2013;30&#xb0;C for 30 min; using the Blank solution as reference, measure the absorbance at 700 nm, and calculate the Total phosphorus content of the sample using <xref ref-type="disp-formula" rid="eq1">Equation 1</xref>:</p>
<disp-formula id="eq1">
<label>(1)</label>
<mml:math display="block" id="M1">
<mml:mrow>
<mml:mi>L</mml:mi>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mo>%</mml:mo>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>C</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>V</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>D</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mn>10</mml:mn>
</mml:mrow>
<mml:mn>4</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<p>where C is the measured concentration (mg/L), V is the final volume (100 mL), D is the dilution factor, and m is the sample weight (0.2000 g)</p>
</sec>
<sec id="s2_5">
<label>2.5</label>
<title>Spectrum data conversion</title>
<p>In the spectral data processing process, specific Spectral transformation can be used to mitigate the effects of environmental factors and interferences, improve the Signal-to-noise ratio, and make the spectral form more suitable for Korla&#xa0;fragrant pear LTP. In this study, several mathematical transformations were applied to the original spectra, generating six types of spectral data: Original absorbance (A), MSC, SG, FD, SD, MSC+FD,MSC+SD, SG+FD, SG+SD (As shown in <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref>). MSC is a Normalization technique that reduces baseline drift, improves the Signal-to-noise ratio, and better reveals differences and similarities among samples; it is commonly used to eliminate scattering effects on spectral data (<xref ref-type="bibr" rid="B16">Gautam et&#xa0;al., 2015</xref>). SG is a Local smoothing method based on Polynomial fitting, which performs Weighted filtering on Spectral data through a Sliding window, achieving Random noise reduction while preserving Peak shape and Spectral details (<xref ref-type="bibr" rid="B31">Liu et&#xa0;al., 2016</xref>). FD can enhance the Resolution of spectra, distinguish Overlapping peaks, and is suitable for determining Peak position and boundary, especially for eliminating Background interference in Quantitative analysis (<xref ref-type="bibr" rid="B42">Sonobe and Hirono, 2023</xref>). SD is more sensitive in identifying Peak inflection points and shoulder peaks, and suppresses the Broad background signal, commonly used in the analysis of fine structures in complex spectra (<xref ref-type="bibr" rid="B15">Gao et&#xa0;al., 2016</xref>). MSC+FD first corrects the Scattering effect using MSC, then applies First-order derivative (FD) to eliminate residual baseline drift and improve peak Resolution, making it suitable for scenarios with strong scattering interference and requiring precise Peak positioning (<xref ref-type="bibr" rid="B47">Wang et&#xa0;al., 2023</xref>). MSC+SD enhances spectral details after MSC reduces scattering, with the second-order derivative (SD) further identifying subtle differences in overlapping peaks, suitable for complex samples requiring resolution of highly overlapping peaks (<xref ref-type="bibr" rid="B51">Xie et&#xa0;al., 2018</xref>). SG+FD first applies SG smoothing to reduce noise, then calculates the first-order derivative (FD), avoiding amplification of noise in derivative results, thus balancing noise suppression and resolution enhancement in spectra with high noise levels (<xref ref-type="bibr" rid="B63">Zhou et&#xa0;al., 2024</xref>). SG+SD first uses SG smoothing to reduce noise, after which the second-order derivative (SD) more accurately reflects spectral curvature, avoiding false peaks caused by noise, making it suitable for spectra requiring fine structure analysis where noise is significant (<xref ref-type="bibr" rid="B3">An et&#xa0;al., 2021</xref>).</p>
<table-wrap id="T1" position="float">
<label>Table&#xa0;1</label>
<caption>
<p>Spectral preprocessing methods.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="left">Full name</th>
<th valign="middle" align="left">Abbreviation</th>
<th valign="middle" align="left">Function</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="left">Multiple Scattering Correction</td>
<td valign="middle" align="left">MSC</td>
<td valign="middle" align="left">Eliminate scattering</td>
</tr>
<tr>
<td valign="middle" align="left">Savitzky-Golay smoothing</td>
<td valign="middle" align="left">SG</td>
<td valign="middle" align="left">Reduce noise</td>
</tr>
<tr>
<td valign="middle" align="left">First derivative</td>
<td valign="middle" align="left">FD</td>
<td valign="middle" align="left">Enhance spectral resolution</td>
</tr>
<tr>
<td valign="middle" align="left">Second derivative</td>
<td valign="middle" align="left">SD</td>
<td valign="middle" align="left">Conduct fine structural analysis</td>
</tr>
<tr>
<td valign="middle" align="left">Multiple Scattering Correction+ First Derivative</td>
<td valign="middle" align="left">MSC+FD</td>
<td valign="middle" align="left">Eliminate scattering + Enhance spectral resolution</td>
</tr>
<tr>
<td valign="middle" align="left">Multiple Scattering Correction+Second Derivative</td>
<td valign="middle" align="left">MSC+SD</td>
<td valign="middle" align="left">Eliminate scattering + Conduct fine structural analysis</td>
</tr>
<tr>
<td valign="middle" align="left">Savitzky-Golay+First Derivative</td>
<td valign="middle" align="left">SG+FD</td>
<td valign="middle" align="left">Reduce noise + Enhance spectral resolution</td>
</tr>
<tr>
<td valign="middle" align="left">Savitzky-Golay+Second Derivative</td>
<td valign="middle" align="left">SG+SD</td>
<td valign="middle" align="left">Reduce noise + Conduct fine structural analysis</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s2_6">
<label>2.6</label>
<title>Extraction of spectrum characteristic bands</title>
<p>To reduce the band redundancy and interference of high-dimensional spectral data, feature bands significantly associated with leaf total phosphorus content (LTP) (LTP) were selected from the spectral data to improve modeling accuracy. In this study, the Competitive adaptive reweighted sampling (CARS) algorithm (<xref ref-type="bibr" rid="B43">Sun et&#xa0;al., 2021</xref>; <xref ref-type="bibr" rid="B52">Xing et&#xa0;al., 2021</xref>) was adopted, based on the principle of darwin&#x2019;s theory of biological evolution&#x2019;s survival of the fittest. Efficient dimensionality reduction of spectral variables was achieved by coupling Partial least squares (PLS) modeling with an adaptive variable screening mechanism. This algorithm evaluates the importance of variables based on the absolute percentage evaluation of PLS model coefficients, generates an initial wavelength subset through Monte Carlo sampling (MCS), and dynamically adjusts variable weight by incorporating an Exponential decay function to strengthen the selective retention of high-contribution bands. Simultaneously employing the Adaptive weighted sampling (ARS) strategy, wavelengths are weighted screening based on the Absolute value of coefficients, prioritizing the retention of Important bands and eliminating Redundant information. After Multi-round iterative optimization, the final Characteristic wavelength combination highly correlated with LTP is obtained, providing an efficient Input variable set for subsequent Regression modeling.</p>
</sec>
<sec id="s2_7">
<label>2.7</label>
<title>Machine learning modeling</title>
<p>Based on the above trait selection results, three algorithms&#x2014;Random forest (RF), Support vector machine (SVR), and Back Propagation(BP)neural network&#x2014;were used to construct a Korla Fragrant Pear LTP estimation model.</p>
<p>RF (<xref ref-type="bibr" rid="B18">Guo and Hao, 2021</xref>)reduces model variance by integrating multiple decision trees, with core parameter settings as follows: Number of decision trees (n _ estimators) was set to four gradients&#x2014;100, 250, 500, and 750&#x2014;where the smaller value (100) was used to explore the model baseline performance, medium values (250, 500) balanced computational efficiency and ensemble effect, and the larger value (750) verified the fitting ability under extreme ensemble scale (<xref ref-type="bibr" rid="B35">Rhodes et&#xa0;al., 2023</xref>); min_samples_leaf (min_samples_leaf) was set to five gradients&#x2014;1, 3, 5, 7, and 10&#x2014;where the minimum value (1) allowed the decision tree to grow sufficiently to capture subtle variations, medium values (3, 5) suppressed noise dissociation, and larger values (7, 10) enforced simplification of tree structure to reduce complexity (<xref ref-type="bibr" rid="B25">Jeong et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B37">Scornet, 2016</xref>).</p>
<p>SVR (<xref ref-type="bibr" rid="B46">Virnodkar et&#xa0;al., 2020</xref>)employed four types of kernel functions for comparison: Radial basis kernel function (RBF), Linear kernel function, Polynomial kernel function, and Sigmoid kernel function. The optimal parameter values of the penalty coefficient (C) and the kernel function parameter (&#x3b3;) were determined through grid search optimization, thereby identifying the optimum parameter and the kernel function.</p>
<p>BP neural network (<xref ref-type="bibr" rid="B29">Li et&#xa0;al., 2021</xref>)consists of input, output, and an intermediate hidden layer. A single hidden layer is used, with the number of nodes set to 1&#x2013;10. The model iterates 1000 times, with a learning rate of 0.01 and a training target error of 1&#xd7;10<sup>-</sup>&#x2076;. Six training functions (trainlm, traingd, trainscg, traingdx, trainbfg, and traincgb) are compared to determine the optimal parameters and the best training function (<xref ref-type="bibr" rid="B49">Wang et&#xa0;al., 2023</xref>).</p>
<p>This paper conducts comprehensive comparative experiments with the currently recognized advanced baseline model in the field. For this, we selected three high-performing and widely used Representative model for spectral analysis as advanced representatives of the baseline model. Partial least squares regression (PLSR) and Light gradient boosting machine (LightGBM) and One-dimensional convolutional neural network (1D-CNN).</p>
<p>PLSR is a Regression modeling method proposed in the early 1980s by Svante Wold and others for handling High-dimensional, multicollinear data (<xref ref-type="bibr" rid="B27">Li et&#xa0;al., 2024</xref>). It achieves prediction of Y by extracting Latent Variables (Latent Variables) from the Independent variable (X) and Dependent variable (Y), and establishing a linear regression relationship between these latent variables. It is particularly suitable for the Chemometrics field such as Spectroscopy and Chromatography. LightGBM was proposed in 2017 by Microsoft Research Asia, and is an efficient implementation of the Gradient Boosting Decision Tree (GBDT) framework, greatly improving Training speed and reducing Memory consumption, while maintaining high prediction accuracy, making it perform exceptionally well on Large-scale data (<xref ref-type="bibr" rid="B21">Gupta et&#xa0;al., 2021</xref>).1D-CNN originated from the LeNet-5 architecture proposed by Yann LeCun and others in the late 1980s and early 1990s for Handwritten digit recognition (designed for 2D images), which was the prototype of the Convolutional Neural Network (CNN). 1D-CNN can automatically learn Local patterns and Multi-scale features in data, making it highly suitable for processing one-dimensional signal such as time series analysis, audio frequency processing, and Near-infrared spectroscopy (<xref ref-type="bibr" rid="B32">Liu et&#xa0;al., 2022</xref>).</p>
</sec>
<sec id="s2_8">
<label>2.8</label>
<title>Medel evaluation methods</title>
<p>This study implemented the aforementioned Regression algorithm using MATLAB R2024b, and comprehensively evaluated Model performance using three metrics: Coefficient of determination (R&#xb2;), Root mean square error (RMSE), and Residual prediction deviation (RPD):</p>
<p>R&#xb2;: Measures Model goodness of fit. The value ranges from 0 to 1. The closer R&#xb2; is to 1, the higher the agreement between the Model predicted value and the measured value (<xref ref-type="bibr" rid="B2">Ahmad Yasmin et&#xa0;al., 2021</xref>). The calculation formula is shown in <xref ref-type="disp-formula" rid="eq2">Equation 2</xref>.</p>
<p>RMSE: Quantifies the absolute magnitude of Prediction error. The smaller the RMSE value, the higher the Model prediction accuracy (<xref ref-type="bibr" rid="B26">Li et&#xa0;al., 2024</xref>). The calculation formula is shown in <xref ref-type="disp-formula" rid="eq3">Equation 3</xref>.</p>
<p>RPD: Reflects the predictive capability of the model. Evaluation criteria: RPD &gt; 3: Excellent Model prediction ability; 2&lt; RPD &#x2264; 3: The model can be used for preliminary prediction; RPD &#x2264; 2: Poor Model prediction ability. The calculation formula is shown in <xref ref-type="disp-formula" rid="eq4">Equation 4</xref>.In model evaluation, the Dataset is divided into training sets and Test set at a ratio of 3:1, and the above metrics are calculated separately to comprehensively assess the model&#x2019;s Fitting effect, Prediction accuracy, and Generalization ability (<xref ref-type="bibr" rid="B10">Chu et&#xa0;al., 2023</xref>).</p>
<disp-formula id="eq2">
<label>(2)</label>
<mml:math display="block" id="M2">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>m</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>p</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>m</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq3">
<label>(3)</label>
<mml:math display="block" id="M3">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>E</mml:mi>
<mml:mo>=</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mi>n</mml:mi>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:munderover>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>m</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>p</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq4">
<label>(4)</label>
<mml:math display="block" id="M4">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>P</mml:mi>
<mml:mi>D</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi>S</mml:mi>
<mml:mi>y</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>E</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mi>m</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>E</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<p>the sample size; <inline-formula>
<mml:math display="inline" id="im1">
<mml:mrow>
<mml:msub>
<mml:mtext>y</mml:mtext>
<mml:mtext>m</mml:mtext>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula>
<mml:math display="inline" id="im2">
<mml:mrow>
<mml:msub>
<mml:mtext>y</mml:mtext>
<mml:mtext>p</mml:mtext>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>They are the actual value and the predicted value of Korla fragrant pear Leaf Total Potassium respectively; <inline-formula>
<mml:math display="inline" id="im3">
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:math>
</inline-formula>It is the average value of the actual Korla fragrant pear Leaf Total Potassium; <inline-formula>
<mml:math display="inline" id="im4">
<mml:mrow>
<mml:msub>
<mml:mtext>S</mml:mtext>
<mml:mi>y</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>Is the standard deviation of the Leaf Total Potassium measurement value of Korla fragrant pear</p>
</sec>
</sec>
<sec id="s3" sec-type="results">
<label>3</label>
<title>Results and analysis</title>
<sec id="s3_1">
<label>3.1</label>
<title>Content of korla fragrant pear total phosphorus in leaves at different periods</title>
<p>As shown in <xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref> Box plot, the LTP content in Korla Fragrant Pear leaves exhibited significant distribution differences across various growing stage, providing a basis for Stage-based modeling. Samples from the fruit-setting period generally had lower overall content but included individual high-value outliers, reflecting periodic fluctuations in Phosphorus demand during this stage. This distribution facilitates the model&#x2019;s ability to capture differences in Spectral response under low and high phosphorus conditions. The distribution of LTP content during the fruit swelling period was relatively concentrated with low dispersion, indicating more stable Phosphorus levels at this stage, allowing modeling to focus on identifying representative Spectral characteristics. The maturity period showed more high-value samples and extreme high values, which, while increasing the difficulty of model identification, also provided critical sample support for establishing Quantitative prediction within the high-content range. The above distribution characteristics indicate that there are significant differences in the Growth period LTP content and their degree of variation among various periods. Therefore, adopting a unified prediction model is unlikely to achieve global optimization. Instead, it is necessary to construct specificity model based on the data characteristics of each period to improve Prediction accuracy and robustness.</p>
</sec>
<sec id="s3_2">
<label>3.2</label>
<title>Analysis of leaf spectral data in korla fragrant pear</title>
<p>This study is based on the use of Near-Infrared Spectroscopy (4000&#x2013;10000 cm<sup>-</sup>&#xb9;) to analyze the Leaf total phosphorus (LTP) content of Korla fragrant pear. Differences in leaf LTP content exist during Different growth stages, manifested as Fruit-setting period&lt; Fruit expansion period&lt; Maturity period, providing a sample basis for Stage-based modeling. It provides a sample foundation for Spectral modeling.</p>
<p>Spectral response indicates that changes in LTP content are significantly associated with the Vibrational absorption of Hydrogen-containing group (such as P-O-H). In the 4000&#x2013;5500 cm<sup>-</sup>&#xb9; range, the Hydrogen group overtone absorption of Phosphorus substance overlaps with that of moisture and Carbohydrates, forming the Core sensitive region for LTP; the Combination frequency and overtone in the 5500&#x2013;7500 cm<sup>-</sup>&#xb9; range can synergistically and complementarily verify differences in total phosphorus (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3A</bold>
</xref>).The line discretization and polymerization of Spectrum curves from different samples reflect the group differences in LTP content. Combined with the Growth period classification of Spectrum plots, distinct dispersion and polymerization of curves in characteristic intervals across different periods are observed (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3B</bold>
</xref>), further providing a basis for establishing the Stage-based prediction model.</p>
</sec>
<sec id="s3_3">
<label>3.3</label>
<title>Spectral data preprocessing</title>
<p>This study applied MSC, SG, FD, SD and their combined methods to preprocess the near-infrared spectra of Korla fragrant pear leaves, aiming to enhance the Spectral characteristics associated with LTP content while reducing noise and scattering interference.</p>
<p>The Original spectrum was influenced by baseline effects and noise, leading to significant signal overlap and indistinct variations (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3A</bold>
</xref>). MSC effectively removed baseline drift caused by particle size and surface scattering, markedly improving spectral consistency and comparability (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4A</bold>
</xref>); SG suppressed random noise while preserving the original Peak shape, thereby increasing the Signal-to-noise ratio and help to highlight phosphorus-related absorption features (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4B</bold>
</xref>); FD processing amplified the dynamic differences of absorption peak related to LTP content by calculating the spectral rate of change, which improved feature discernibility (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4C</bold>
</xref>); SD further accentuated subtle changes in Spectral curvature, proving particularly useful for extracting weak signals from low-content samples (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4D</bold>
</xref>).Combined preprocessing strategy integrates the advantages of individual methods, further enhancing Spectrum quality. MSC+FD eliminates physical scattering while amplifying dynamic spectral features, making it more effective at capturing variations in LTP content (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4E</bold>
</xref>); MSC+SD improves detail resolution on the basis of scatter correction, supplying the model with more stable and refined input (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4F</bold>
</xref>); SG+FD preserves and accentuates spectral changes induced by chemical constituents while reducing noise, thereby balancing the need for Smooth and feature enhancement (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4G</bold>
</xref>); SG+SD achieves both noise suppression and high-frequency detail enhancement, making it suitable for dynamic monitoring and Modeling of LTP content across the whole growth period (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4H</bold>
</xref>).</p>
<p>The results demonstrate that different preprocessing methods improve Spectrum quality in various aspects, including noise suppression, removal of physical interference, and enhancement of dynamic features. Combined methods exhibit stronger adaptability and synergistic Gain effects. Subsequently, the optimal pretreatment strategy will be selected based on Model performance metrics, laying the groundwork for high-accuracy prediction of LTP content.</p>
</sec>
<sec id="s3_4">
<label>3.4</label>
<title>Spectral data and correlation analysis with LTP</title>
<p>This study evaluated the optimization effects of various methods on Phosphorus information extraction by analyzing the Correlation (r) between different Preprocessed spectra and LTP content. The main findings are summarized below (representative r values are provided in <xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5</bold>
</xref>):</p>
<fig id="f5" position="float">
<label>Figure&#xa0;5</label>
<caption>
<p>
<bold>(A)</bold> Original spectrum; <bold>(B)</bold> MSC; <bold>(C)</bold> SG; <bold>(D)</bold> FD; <bold>(E)</bold> SD; <bold>(F)</bold> MSC+FD; <bold>(G)</bold> MSC+SD; <bold>(H)</bold> SG+FD; <bold>(I)</bold> SG+SD.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g005.tif">
<alt-text content-type="machine-generated">Nine graphs labeled A to I display spectral data with various preprocessing methods. Each plot shows a spectral curve against wavenumber (cm&#x207b;&#xb9;) on the x-axis, ranging from 4000 to 10000 cm&#x207b;&#xb9;, and an axis labeled 'r' ranging approximately from -1 to 1. Methods include A (raw), MSC, SG, FD, SD, MSC+FD, MSC+SD, SG+FD, and SG+SD, each altering the spectral shape differently.</alt-text>
</graphic>
</fig>
<p>The Original spectrum (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5A</bold>
</xref>) displays broad and smooth Peak shape, rendering it susceptible to scattering and Noise interference; Multiplicative scatter correction processing (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>) effectively suppressed physical interference; Savitzky-Golay processing (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5C</bold>
</xref>) showed no noticeable improvement; First Derivative processing (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5D</bold>
</xref>) amplified both Dynamic correlation details and Noise interference; Second Derivative processing (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5E</bold>
</xref>) improved the ability to extract trace phosphorus-containing components and detect Weak correlation with LTP.</p>
<p>Combined preprocessing methods exhibited stronger Synergistic effect. The MSC+FD approach (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5F</bold>
</xref>) significantly enhanced the Recognition ability of LTP in Highly discrete samples; MSC+ SD (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5G</bold>
</xref>) improved the resolution of LTP information in Low phosphorus content samples; SG+ FD(<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5H</bold>
</xref>) strengthened the correlation between spectral data and LTP while enhancing Interpretability; SG+ SD (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5I</bold>
</xref>) reduced Noise interference and accentuated Weak absorption difference, thereby improving the Dynamic monitoring ability of LTP during Different growth stages. In summary, the Combined pre-processing method can more effectively extract Spectral characteristics associated with LTP content, thereby providing a more reliable data foundation for subsequent Modeling.</p>
<p>Correlation Analysis between Different Spectral Data and LTP in <xref ref-type="fig" rid="f5">
<bold>Figure 5</bold>
</xref>: (A) Original spectrum; (B) MSC; (C) SG; (D) FD; (E)&#xa0;SD; (F) MSC+FD; (G) MSC+SD; (H) SG+FD; (I) SG+SD.</p>
</sec>
<sec id="s3_5">
<label>3.5</label>
<title>Selection of LTP characteristic bands in korla fragrant pear</title>
<p>This study utilized the Competitive Adaptive Reweighted Sampling algorithm to screening feature bands highly correlated with LTP content from preprocessed near-infrared spectra, and analyzed the distribution characteristics of these bands across both the whole growth period and Different growth stages (as shown in <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>).</p>
<fig id="f6" position="float">
<label>Figure&#xa0;6</label>
<caption>
<p>Selection of korla fragrant pear LTP characteristic bands: <bold>(A)</bold> represents the selection for the spectral data of the whole growth period; <bold>(B)</bold> represents the selection for the spectral data of the fruit bearing periods; <bold>(C)</bold> represents the selection for the spectral data of the fruit swelling period; <bold>(D)</bold> represents the selection for the spectral data of the fruit ripening period.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g006.tif">
<alt-text content-type="machine-generated">Four plots labeled A, B, C, and D display different configurations of CARS data versus wavenumber in wavenumbers per centimeter. Each plot has eight data series: SG+SD-CARS, SG+FD-CARS, MSC+SD-CARS, MSC+FD-CARS, SD-CARS, FD-CARS, SG-CARS, and MSC-CARS, shown in distinct colors. The graphs are on a light blue background with wavenumber ranging from four thousand to ten thousand wavenumbers per centimeter on the x-axis. Each configuration is plotted separately, highlighting variations across different methods.</alt-text>
</graphic>
</fig>
<p>The results demonstrate that combined preprocessing methods (such as MSC+FD,SG+SD) consistently extracted denser and more comprehensive feature bands across all Different growth stages, proving particularly effective at capturing subtle chemical absorption variations. In contrast, single preprocessing methods (e.g., Multiplicative scatter correction or Second Derivative) mainly focused on core absorption peak regions, resulting in fewer but more specific feature bands.</p>
<p>Regarding growth stage differences: during the fruit-setting stage, LTP content was high and exhibited high variability. Combined preprocessing methods produced densely clustered feature bands in certain spectral regions, effectively adapting to high-phosphorus absorption variations and capturing dynamic differential features. During the fruit-expansion stage, LTP content remained stable, and feature bands were more uniformly distributed across the 4000&#x2013;8000 cm<sup>-</sup>&#xb9; range. Combined methods formed continuous characteristic zones in some bands, aligning with chemical equilibrium states during this stable period and enabling the identification of more comprehensively correlated bands. In the maturity stage, LTP content was low and absorption signals were weak, making it necessary to rely on combined preprocessing to enhance the extraction of bands sensitive to trace components. The Feature band selection results provide critical input for building subsequent staged LTP content prediction models. By leveraging Growth period characteristics, suitable Pretreatment strategies can be selected to improve model accuracy and specificity.</p>
</sec>
<sec id="s3_6">
<label>3.6</label>
<title>Korla fragrant pear LTP estimation model</title>
<p>This study constructed Growth-period-specific models and an intertemporal general model based on the LTP content and Spectral data of Korla fragrant pear leaves at Different growth stages, using Coefficient of determination (R&#xb2;), Root mean square error (RMSE), and Residual prediction deviation (RPD) as evaluation metrics for Model performance (Supplementary Material 1). The results indicated that each Growth-period-specific model significantly outperformed the intertemporal general model in predicting LTP content.</p>
<p>As shown in <xref ref-type="table" rid="T1">
<bold>Tables&#xa0;1</bold>
</xref> and <xref ref-type="table" rid="T2">
<bold>2</bold>
</xref>, the optimal model for the fruit-setting period (FD+CARS-BP) achieved R&#xb2; = 0.89, RMSE = 0.0212, RPD = 3.1696 on the Training set and R&#xb2; = 0.88, RMSE = 0.0241, RPD = 2.6963 on the validation set, demonstrating a stronger ability to capture the dynamic Spectral characteristics of highly discrete LTP content. Its performance was significantly superior to that of the MSC-CARS-BP general model. The optimal model for the Fruit expansion period (SG+FD-CARS-BP) achieved Coefficient of determination (R&#xb2;) values of 0.86 and 0.83 on the Training set and validation set, respectively, with Root mean square error (RMSE) values of 0.0211 and 0.0254, and RPD values of 2.6721 and 2.4571, demonstrating better adaptation to the relatively stable Spectrum-chemical state during this period. The optimal model for the Maturity period (SG+SD-CARS-BP) achieved Coefficient of determination (R&#xb2;) values above 0.85 on both training and validation sets, with Root mean square error (RMSE) below 0.021 and RPD exceeding 2.68, indicating effective resolution of weak Spectrum signals from low-content LTP and significantly outperforming the general model.</p>
<table-wrap id="T2" position="float">
<label>Table&#xa0;2</label>
<caption>
<p>Indicators of the best models under each machine learning algorithm in each period.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" rowspan="2" align="center">Period</th>
<th valign="middle" rowspan="2" align="center">Model</th>
<th valign="middle" colspan="3" align="center">Training set</th>
<th valign="middle" colspan="3" align="center">Validation set</th>
</tr>
<tr>
<th valign="middle" align="center">R<sup>2</sup>
</th>
<th valign="middle" align="center">RMSE</th>
<th valign="middle" align="center">RPD</th>
<th valign="middle" align="center">R<sup>2</sup>
</th>
<th valign="middle" align="center">RMSE</th>
<th valign="middle" align="center">RPD</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" rowspan="3" align="center">Entire growth period</td>
<td valign="middle" align="center">MSC+SD-CARS-RF</td>
<td valign="middle" align="center">0.80</td>
<td valign="middle" align="center">0.0269</td>
<td valign="middle" align="center">2.2094</td>
<td valign="middle" align="center">0.76</td>
<td valign="middle" align="center">0.0334</td>
<td valign="middle" align="center">2.0521</td>
</tr>
<tr>
<td valign="middle" align="center">MSC-CARS-BP</td>
<td valign="middle" align="center">0.84</td>
<td valign="middle" align="center">0.0258</td>
<td valign="middle" align="center">2.4436</td>
<td valign="middle" align="center">0.81</td>
<td valign="middle" align="center">0.0268</td>
<td valign="middle" align="center">2.3764</td>
</tr>
<tr>
<td valign="middle" align="center">MSC+FD-CARS-SVM</td>
<td valign="middle" align="center">0.79</td>
<td valign="middle" align="center">0.0287</td>
<td valign="middle" align="center">2.1862</td>
<td valign="middle" align="center">0.78</td>
<td valign="middle" align="center">0.0290</td>
<td valign="middle" align="center">2.0917</td>
</tr>
<tr>
<td valign="middle" rowspan="3" align="center">Fruit setting period</td>
<td valign="middle" align="center">SG+FD-CARS-RF</td>
<td valign="middle" align="center">0.85</td>
<td valign="middle" align="center">0.0267</td>
<td valign="middle" align="center">2.5807</td>
<td valign="middle" align="center">0.77</td>
<td valign="middle" align="center">0.0295</td>
<td valign="middle" align="center">2.0965</td>
</tr>
<tr>
<td valign="middle" align="center">FD+CARS-BP</td>
<td valign="middle" align="center">0.89</td>
<td valign="middle" align="center">0.0212</td>
<td valign="middle" align="center">3.1696</td>
<td valign="middle" align="center">0.88</td>
<td valign="middle" align="center">0.0241</td>
<td valign="middle" align="center">2.9663</td>
</tr>
<tr>
<td valign="middle" align="center">MSC+SD-CARS-SVM</td>
<td valign="middle" align="center">0.84</td>
<td valign="middle" align="center">0.0279</td>
<td valign="middle" align="center">2.2024</td>
<td valign="middle" align="center">0.79</td>
<td valign="middle" align="center">0.0281</td>
<td valign="middle" align="center">1.9787</td>
</tr>
<tr>
<td valign="middle" rowspan="3" align="center">Fruit Enlargement Stage</td>
<td valign="middle" align="center">MSC+SD-CARS-RF</td>
<td valign="middle" align="center">0.80</td>
<td valign="middle" align="center">0.0257</td>
<td valign="middle" align="center">2.2293</td>
<td valign="middle" align="center">0.82</td>
<td valign="middle" align="center">0.0238</td>
<td valign="middle" align="center">2.3391</td>
</tr>
<tr>
<td valign="middle" align="center">SG+FD-CARS-BP</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0211</td>
<td valign="middle" align="center">2.6721</td>
<td valign="middle" align="center">0.83</td>
<td valign="middle" align="center">0.0254</td>
<td valign="middle" align="center">2.4571</td>
</tr>
<tr>
<td valign="middle" align="center">MSC-CARS-SVM</td>
<td valign="middle" align="center">0.90</td>
<td valign="middle" align="center">0.0194</td>
<td valign="middle" align="center">4.3079</td>
<td valign="middle" align="center">0.81</td>
<td valign="middle" align="center">0.0228</td>
<td valign="middle" align="center">2.2185</td>
</tr>
<tr>
<td valign="middle" rowspan="3" align="center">Fruit Ripening Stage</td>
<td valign="middle" align="center">SG+FD-CARS-RF</td>
<td valign="middle" align="center">0.81</td>
<td valign="middle" align="center">0.0231</td>
<td valign="middle" align="center">2.3696</td>
<td valign="middle" align="center">0.80</td>
<td valign="middle" align="center">0.0237</td>
<td valign="middle" align="center">2.2457</td>
</tr>
<tr>
<td valign="middle" align="center">SG+SD-CARS-BP</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0203</td>
<td valign="middle" align="center">2.6977</td>
<td valign="middle" align="center">0.85</td>
<td valign="middle" align="center">0.0207</td>
<td valign="middle" align="center">2.6844</td>
</tr>
<tr>
<td valign="middle" align="center">MSC-CARS-SVM</td>
<td valign="middle" align="center">0.82</td>
<td valign="middle" align="center">0.0227</td>
<td valign="middle" align="center">3.0751</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0218</td>
<td valign="middle" align="center">2.2088</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>Linear fitting results (<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>) further support the above conclusions. The Coefficient of determination (R&#xb2;) values for the Fruit-setting period model reached 0.905 and 0.888 on the training and validation sets, respectively, indicating its strong fitting and Generalization ability even with highly Variant data. The Coefficient of determination (R&#xb2;) values for the Fruit expansion period model were all above 0.83, and those for the Maturity period model exceeded 0.86, further confirming the adaptability and stability of the specificity model across Different growth stages. Studies have shown that the Growth-period-specific model, by aligning with the differences in LTP content and Spectral characteristics across various periods, significantly improves prediction accuracy, providing a reliable model option for precise phosphorus nutrition management in orchards. Therefore, it is recommended to use the FD+CARS-BP model during the fruit-setting period, the SG+FD-CARS-BP model during the fruit swelling period, and the SG+SD-CARS-BP model during the maturity period for predicting the Korla fragrant pear LTP content.</p>
<fig id="f7" position="float">
<label>Figure&#xa0;7</label>
<caption>
<p>
<bold>(A)</bold> shows the linear fit between the measured and predicted values of the training sets for the fruit bearing periodsFD+CARS-BP model; <bold>(B)</bold> shows the linear fit between the measured and predicted values of the validation set for the fruit bearing periodsFD+CARS-BP model; <bold>(C)</bold> shows the linear fit between the measured and predicted values of the training sets for the fruit swelling periodSG+FD-CARS-BP model; <bold>(D)</bold> shows the linear fit between the measured and predicted values of the validation set for the fruit swelling periodSG+FD-CARS-BP model; <bold>(E)</bold> shows the linear fit between the measured and predicted values of the training sets for the fruit ripening periodSG+SD-CARS-BP model; <bold>(F)</bold> shows the linear fit between the measured and predicted values of the validation set for the fruit ripening periodSG+SD-CARS-BP model.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g007.tif">
<alt-text content-type="machine-generated">Scatter plots labeled A to F showing predicted versus measured values with trend lines. Each plot displays a positive correlation with R-squared values as follows: A - 0.90504, B - 0.88785, C - 0.86243, D - 0.83488, E - 0.87246, F - 0.86146. The axes represent percentages of predicted and measured values.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3_7">
<label>3.7</label>
<title>Model parameters and function selection</title>
<p>In the Random forest (RF) Modeling process, the number of Decision trees and the value of min_samples_leaf are key hyper parameters that directly affect model complexity and Generalization ability. With too few trees, the model&#x2019;s fitting ability is inadequate, leading to poor fitting. As the number of trees increases, the model improves prediction stability through ensemble learning. However, beyond a certain point, the performance Gain becomes marginal, while computational cost and Overfitting risk rise (<xref ref-type="bibr" rid="B24">Huang et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B19">Guo et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B11">Dabiri et&#xa0;al., 2022</xref>). Taking the fruit ripening period SG+FD-CARS-RF model as an example (<xref ref-type="fig" rid="f8">
<bold>Figures&#xa0;8A&#x2013;D</bold>
</xref>), when the number of trees is 500 and min_samples_leaf is 5, the difference in Coefficient of determination (R&#xb2;) between the Training set and the validation set is the smallest (0.0112), and the Root mean square error (RMSE) difference is only 0.006&#x2014;significantly better than other parameter combinations. This indicates that this configuration maintains strong Generalization ability while mitigating overfitting, and was therefore identified as the optimal parameter set. For the SVM model, using the mature stage Multiplicative scatter correction-Competitive Adaptive Reweighted Sampling-SVM as an example, it is essential to optimize the regularization parameter C and the kernel parameter &#x3b3;. As illustrated in <xref ref-type="fig" rid="f9">
<bold>Figures&#xa0;9A&#x2013;D</bold>
</xref>, when C = 5 and &#x3b3;=0.1, the model performs well on both the Training set and the validation set (with Coefficient of determination (R&#xb2;) values of 0.8185 and 0.8552, and Root mean square error (RMSE) values of 0.0227 and 0.0218, respectively), demonstrating a good balance. Although a higher Training set Coefficient of determination (R&#xb2;) of 0.8716 is achieved when C = 10 and &#x3b3;=0.3, the validation set Coefficient of determination (R&#xb2;) drops significantly to 0.6943, indicating clear overfitting. Thus, C = 5 and &#x3b3;=0.1 are identified as the optimal parameters. Further comparison among different kernel functions shows that the Radial basis kernel function (RBF) delivers the best overall performance, with the smallest discrepancy in Coefficient of determination (R&#xb2;) between the validation set and the Training set. In contrast, although the Polynomial kernel function performs well on the Training set, its validation set Coefficient of determination (R&#xb2;) decreases by up to 0.3792, reflecting inadequate Generalization ability. This suggests that the RBF kernel is more suitable for the characteristics of the present dataset. In the BP neural network model, using the fruit-setting period FD+CARS-BP as an example, the number of nodes in the hidden layer significantly influences the model&#x2019;s expressive power. As shown in <xref ref-type="fig" rid="f10">
<bold>Figures&#xa0;10A, B</bold>
</xref>, when the hidden layer contains 5 nodes, both the Training set and validation set show high Coefficient of determination (R&#xb2;) and low Root mean square error (RMSE), indicating that this configuration maintains strong fitting ability without noticeable overfitting. Further comparison among different training functions shows that the trainlm function performs best in this model, achieving a validation set Coefficient of determination (R&#xb2;) of 0.88 and an Root mean square error (RMSE) of 0.0241, surpassing other training functions and demonstrating its superior suitability for the given data structure and task complexity.</p>
<fig id="f8" position="float">
<label>Figure&#xa0;8</label>
<caption>
<p>RFhyper parameter settings: <bold>(A)</bold> is R2 of training sets; <bold>(B)</bold> is Root mean square error (RMSE) of training sets; <bold>(C)</bold> is R2 of the validation set; <bold>(D)</bold> is Root mean square error (RMSE) of the validation set.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g008.tif">
<alt-text content-type="machine-generated">Four heatmaps labeled A, B, C, and D compare the impact of the number of decision trees and minimum leaf node samples on R-squared or RMSE values. A and B show R-squared values ranging from 0.7090 to 0.9670, with varying color intensities. C and D depict RMSE values from 0.009550 to 0.03650, also with changing colors. The x-axis represents the number of decision trees (100, 250, 500, 750), and the y-axis shows the minimum number of leaf node samples (1, 3, 5, 7, 10).</alt-text>
</graphic>
</fig>
<fig id="f9" position="float">
<label>Figure&#xa0;9</label>
<caption>
<p>SVM hyper parameter settings and kernel function selection: <bold>(A)</bold> is R2 of training sets; <bold>(B)</bold> is Root mean square error (RMSE) of training sets; <bold>(C)</bold> is R2 of validation set; <bold>(D)</bold> is Root mean square error (RMSE) of validation set; <bold>(E)</bold> is R2 of four kernel functions for training sets and validation set; <bold>(F)</bold> is Root mean square error (RMSE) of four kernel functions for training sets and validation set.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g009.tif">
<alt-text content-type="machine-generated">Heatmaps and bar charts analyze model performance with different kernel functions. Panels A and B display R&#xb2; values; C and D show RMSE values based on regularization and gamma parameters. Panels E and F compare R&#xb2; and RMSE across training and verification sets for various kernel functions.</alt-text>
</graphic>
</fig>
<fig id="f10" position="float">
<label>Figure&#xa0;10</label>
<caption>
<p>BP neural networkhyper parameter settings and kernel function selection: <bold>(A)</bold> is R2 and Root mean square error (RMSE) of training sets; <bold>(B)</bold> is R2 and Root mean square error (RMSE) of the validation set; <bold>(C)</bold> is R2 of six training functions; <bold>(D)</bold> is Root mean square error (RMSE) of six training functions.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1666460-g010.tif">
<alt-text content-type="machine-generated">Graphs A and B compare R&#xb2; and RMSE values against the number of hidden layer nodes, showing variable trends. Graphs C and D present bar charts of R&#xb2; and RMSE for different training algorithms, comparing training and verification sets.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3_8">
<label>3.8</label>
<title>Performance comparison with advanced baseline model</title>
<p>To evaluate the performance of the optimal models selected for each Growth period, this study conducted a comprehensive comparison with widely recognized advanced baseline models in the field, namely PLSR, 1D-CNN, and LightGBM. Each baseline model employed the same Pretreatment and Feature band selection methods as the corresponding optimal model for the respective growth stage (Fruit-setting period: First Derivative-Competitive Adaptive Reweighted Sampling; Fruit-expanding period:SG+FD-CARS; Maturity period:SG+SD-CARS) to ensure a fair comparison (<xref ref-type="table" rid="T3"><bold>Table 3</bold></xref>).</p>
<table-wrap id="T3" position="float">
<label>Table&#xa0;3</label>
<caption>
<p>Performance comparison with advanced baseline model.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" rowspan="2" align="center">Period</th>
<th valign="middle" rowspan="2" align="center">Model</th>
<th valign="middle" colspan="3" align="center">Training set</th>
<th valign="middle" colspan="3" align="center">Validation set</th>
</tr>
<tr>
<th valign="middle" align="center">R&#xb2;</th>
<th valign="middle" align="center">RMSE</th>
<th valign="middle" align="center">RPD</th>
<th valign="middle" align="center">R&#xb2;</th>
<th valign="middle" align="center">RMSE</th>
<th valign="middle" align="center">RPD</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" rowspan="4" align="center">Fruit setting period</td>
<td valign="middle" align="center">BP</td>
<td valign="middle" align="center">0.89</td>
<td valign="middle" align="center">0.0212</td>
<td valign="middle" align="center">3.17</td>
<td valign="middle" align="center">0.88</td>
<td valign="middle" align="center">0.0241</td>
<td valign="middle" align="center">2.97</td>
</tr>
<tr>
<td valign="middle" align="center">PLSR</td>
<td valign="middle" align="center">0.83</td>
<td valign="middle" align="center">0.0265</td>
<td valign="middle" align="center">2.54</td>
<td valign="middle" align="center">0.81</td>
<td valign="middle" align="center">0.0305</td>
<td valign="middle" align="center">2.35</td>
</tr>
<tr>
<td valign="middle" align="center">LightGBM</td>
<td valign="middle" align="center">0.93</td>
<td valign="middle" align="center">0.017</td>
<td valign="middle" align="center">3.95</td>
<td valign="middle" align="center">0.85</td>
<td valign="middle" align="center">0.0262</td>
<td valign="middle" align="center">2.74</td>
</tr>
<tr>
<td valign="middle" align="center">1D-CNN</td>
<td valign="middle" align="center">0.91</td>
<td valign="middle" align="center">0.0191</td>
<td valign="middle" align="center">3.52</td>
<td valign="middle" align="center">0.84</td>
<td valign="middle" align="center">0.0268</td>
<td valign="middle" align="center">2.68</td>
</tr>
<tr>
<td valign="middle" rowspan="4" align="center">Fruit Enlargement Stage</td>
<td valign="middle" align="center">BP</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0211</td>
<td valign="middle" align="center">2.67</td>
<td valign="middle" align="center">0.83</td>
<td valign="middle" align="center">0.0254</td>
<td valign="middle" align="center">2.46</td>
</tr>
<tr>
<td valign="middle" align="center">PLSR</td>
<td valign="middle" align="center">0.81</td>
<td valign="middle" align="center">0.024</td>
<td valign="middle" align="center">2.35</td>
<td valign="middle" align="center">0.78</td>
<td valign="middle" align="center">0.028</td>
<td valign="middle" align="center">2.23</td>
</tr>
<tr>
<td valign="middle" align="center">LightGBM</td>
<td valign="middle" align="center">0.90</td>
<td valign="middle" align="center">0.018</td>
<td valign="middle" align="center">3.13</td>
<td valign="middle" align="center">0.82</td>
<td valign="middle" align="center">0.026</td>
<td valign="middle" align="center">2.4</td>
</tr>
<tr>
<td valign="middle" align="center">1D-CNN</td>
<td valign="middle" align="center">0.88</td>
<td valign="middle" align="center">0.0198</td>
<td valign="middle" align="center">2.85</td>
<td valign="middle" align="center">0.81</td>
<td valign="middle" align="center">0.0265</td>
<td valign="middle" align="center">2.36</td>
</tr>
<tr>
<td valign="middle" rowspan="4" align="center">Fruit Ripening Stage</td>
<td valign="middle" align="center">BP</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0203</td>
<td valign="middle" align="center">2.7</td>
<td valign="middle" align="center">0.85</td>
<td valign="middle" align="center">0.0207</td>
<td valign="middle" align="center">2.68</td>
</tr>
<tr>
<td valign="middle" align="center">PLSR</td>
<td valign="middle" align="center">0.83</td>
<td valign="middle" align="center">0.0221</td>
<td valign="middle" align="center">2.48</td>
<td valign="middle" align="center">0.81</td>
<td valign="middle" align="center">0.0231</td>
<td valign="middle" align="center">2.4</td>
</tr>
<tr>
<td valign="middle" align="center">LightGBM</td>
<td valign="middle" align="center">0.92</td>
<td valign="middle" align="center">0.015</td>
<td valign="middle" align="center">3.66</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0199</td>
<td valign="middle" align="center">2.79</td>
</tr>
<tr>
<td valign="middle" align="center">1D-CNN</td>
<td valign="middle" align="center">0.89</td>
<td valign="middle" align="center">0.0182</td>
<td valign="middle" align="center">3.01</td>
<td valign="middle" align="center">0.80</td>
<td valign="middle" align="center">0.0272</td>
<td valign="middle" align="center">2.41</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>As presented in <xref ref-type="table" rid="T4">
<bold>Table&#xa0;4</bold>
</xref>, during the fruit-setting period, the FD+CARS-BP model attained a Coefficient of determination (R&#xb2;) of 0.88 on the validation set, surpassing PLSR (0.81) and 1D-CNN (0.84). Although its R&#xb2; was marginally lower than that of LightGBM (0.85), the model demonstrated a lower Root mean square error (RMSE) (0.0241) and a higher RPD (2.97), indicating more stable and reliable prediction performance. During the fruit swelling period, the SG+FD-CARS-BP model achieved a validation set Coefficient of determination (R&#xb2;) of 0.83 and an Root mean square error (RMSE) of 0.0254, outperforming both PLSR and 1D-CNN, and performing comparably to LightGBM. Notably, while LightGBM attained a high Coefficient of determination (R&#xb2;) on the Training set (0.90), its performance on the validation set declined significantly (0.82), suggesting potential overfitting. In contrast, the model proposed in this study exhibited more consistent performance across both training and validation sets, indicating superior Generalization ability.</p>
<table-wrap id="T4" position="float">
<label>Table&#xa0;4</label>
<caption>
<p>Comparison of intertemporal models.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" rowspan="2" align="center">Model</th>
<th valign="middle" rowspan="2" align="center">Period</th>
<th valign="middle" colspan="3" align="center">Training set</th>
<th valign="middle" colspan="3" align="center">Validation set</th>
</tr>
<tr>
<th valign="middle" align="center">R<sup>2</sup>
</th>
<th valign="middle" align="center">RMSEC</th>
<th valign="middle" align="center">RPD</th>
<th valign="middle" align="center">R<sup>2</sup>
</th>
<th valign="middle" align="center">RMSEC</th>
<th valign="middle" align="center">RPD</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" rowspan="3" align="center">Reproductive Stage<break/>MSC-CARS-BP Model</td>
<td valign="middle" align="center">fruit setting stage</td>
<td valign="middle" align="center">0.94</td>
<td valign="middle" align="center">0.0174</td>
<td valign="middle" align="center">4.0725</td>
<td valign="middle" align="center">0.73</td>
<td valign="middle" align="center">0.0350</td>
<td valign="middle" align="center">1.9914</td>
</tr>
<tr>
<td valign="middle" align="center">Fruit Enlargement Stage</td>
<td valign="middle" align="center">0.82</td>
<td valign="middle" align="center">0.0248</td>
<td valign="middle" align="center">2.3936</td>
<td valign="middle" align="center">0.77</td>
<td valign="middle" align="center">0.0258</td>
<td valign="middle" align="center">2.3936</td>
</tr>
<tr>
<td valign="middle" align="center">Fruit Ripening Period</td>
<td valign="middle" align="center">0.74</td>
<td valign="middle" align="center">0.0280</td>
<td valign="middle" align="center">2.0821</td>
<td valign="middle" align="center">0.69</td>
<td valign="middle" align="center">0.0286</td>
<td valign="middle" align="center">1.8807</td>
</tr>
<tr>
<td valign="middle" rowspan="3" align="center">Fruit Setting Period<break/>FD+CARS-BP Model</td>
<td valign="middle" align="center">fruit setting period</td>
<td valign="middle" align="center">0.89</td>
<td valign="middle" align="center">0.0212</td>
<td valign="middle" align="center">3.1696</td>
<td valign="middle" align="center">0.88</td>
<td valign="middle" align="center">0.0241</td>
<td valign="middle" align="center">2.9663</td>
</tr>
<tr>
<td valign="middle" align="center">Fruit Enlargement Stage</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0264</td>
<td valign="middle" align="center">2.8344</td>
<td valign="middle" align="center">0.78</td>
<td valign="middle" align="center">0.0323</td>
<td valign="middle" align="center">2.4215</td>
</tr>
<tr>
<td valign="middle" align="center">Fruit Ripening Period</td>
<td valign="middle" align="center">0.82</td>
<td valign="middle" align="center">0.0245</td>
<td valign="middle" align="center">2.3448</td>
<td valign="middle" align="center">0.78</td>
<td valign="middle" align="center">0.0270</td>
<td valign="middle" align="center">2.1481</td>
</tr>
<tr>
<td valign="middle" rowspan="3" align="center">Fruit Swelling Stage<break/>SG+FD-CARS-BP Model</td>
<td valign="middle" align="center">fruit-setting period</td>
<td valign="middle" align="center">0.79</td>
<td valign="middle" align="center">0.0262</td>
<td valign="middle" align="center">2.2430</td>
<td valign="middle" align="center">0.77</td>
<td valign="middle" align="center">0.0285</td>
<td valign="middle" align="center">2.3265</td>
</tr>
<tr>
<td valign="middle" align="center">Fruit Enlargement Stage</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0211</td>
<td valign="middle" align="center">2.6721</td>
<td valign="middle" align="center">0.83</td>
<td valign="middle" align="center">0.0254</td>
<td valign="middle" align="center">2.4571</td>
</tr>
<tr>
<td valign="middle" align="center">Fruit Ripening Period</td>
<td valign="middle" align="center">0.84</td>
<td valign="middle" align="center">0.0224</td>
<td valign="middle" align="center">2.6924</td>
<td valign="middle" align="center">0.76</td>
<td valign="middle" align="center">0.0285</td>
<td valign="middle" align="center">2.1165</td>
</tr>
<tr>
<td valign="middle" rowspan="3" align="center">Fruit Ripening Stage<break/>SG+SD-CARS-BP Model</td>
<td valign="middle" align="center">fruit setting stage</td>
<td valign="middle" align="center">0.80</td>
<td valign="middle" align="center">0.0246</td>
<td valign="middle" align="center">2.2669</td>
<td valign="middle" align="center">0.81</td>
<td valign="middle" align="center">0.0280</td>
<td valign="middle" align="center">2.2801</td>
</tr>
<tr>
<td valign="middle" align="center">Fruit Swelling Stage</td>
<td valign="middle" align="center">0.70</td>
<td valign="middle" align="center">0.0306</td>
<td valign="middle" align="center">1.8600</td>
<td valign="middle" align="center">0.73</td>
<td valign="middle" align="center">0.0330</td>
<td valign="middle" align="center">1.9688</td>
</tr>
<tr>
<td valign="middle" align="center">Fruit Ripening Period</td>
<td valign="middle" align="center">0.86</td>
<td valign="middle" align="center">0.0203</td>
<td valign="middle" align="center">2.6977</td>
<td valign="middle" align="center">0.85</td>
<td valign="middle" align="center">0.0207</td>
<td valign="middle" align="center">2.6844</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>At the Maturity period, the SG+SD-CARS-BP model delivered the best overall predictive performance, with a validation set Coefficient of determination (R&#xb2;) of 0.85, an Root mean square error (RMSE) of 0.0207, and an RPD of 2.68. All metrics surpassed those of PLSR and 1D-CNN. Compared to LightGBM, the proposed model showed better performance in terms of Root mean square error (RMSE) and RPD, further highlighting its accuracy and stability in practical applications. In summary, systematic comparisons with multiple advanced baseline model demonstrate that the Growth period-specific machine learning model developed in this study exhibits consistently excellent and stable predictive ability across different growth stages, confirming the effectiveness and superiority of the Stage-based modeling strategy for monitoring LTP content in Korla fragrant pear leaves. </p>
</sec>
</sec>
<sec id="s4" sec-type="discussion">
<label>4</label>
<title>Discussion</title>
<p>This study systematically investigates the prediction of Leaf total phosphorus (LTP) content in Korla Fragrant Pear, comprehensively revealing the application patterns of Near-Infrared Spectroscopy in fruit nutrition diagnosis through the analysis of Growth period differences, spectral pre-processing refinement, trait screening, and model construction and validation. It provides theoretical and technical support for Precision nutrient management in orchard. Specific discussions are as follows.</p>
<sec id="s4_1">
<label>4.1</label>
<title>The dynamic of LTP content during the fertile period and model adaptability</title>
<p>Different growth stages Korla Fragrant Pear leaf LTP content showed significant differences (P&lt; 0.05), with low and discrete content during the fruit-setting period, stable during the fruit expansion period, and high and discrete during the Maturity period (<xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>). The FD+CARS-BP model performed excellently during the fruit-setting period, with R&#xb2; of training sets reaching 0.90504 and validation set R&#xb2; of 0.88785. FD preprocessing enhanced spectral dynamic differences, and the CARS algorithm accurately screening bands associated with highly discrete LTP, adapting to the &#x201c;high dynamic phosphorus content-complex Spectral response&#x201d; trait (<xref ref-type="bibr" rid="B57">Yu et&#xa0;al., 2024</xref>).;Fruit expansion period SG+FD-CARS-BP model leverages the synergistic effect of SG noise reduction and FD enhancement to balance the &#x201c;Stable Spectra-Basic Phosphorus Absorption&#x201d; relationship, with R2 values of 0.86243 and 0.83488 for training sets and validation set respectively; Maturity period SG+SD-CARS-BP model utilizes SG and SD fine feature extraction to effectively capture trace phosphorus association information, achieving R2 values of 0.87246 and 0.86146 for training sets and validation set respectively. This demonstrates that the Growth-period-specific model significantly improves prediction accuracy by adapting to the dynamics of LTP content across different periods (high dispersion, stable state, low concentration), validating the necessity of &#x201c;Stage-based modeling&#x201d; in fruit nutrition diagnosis. These findings align with the conclusions of Li et&#xa0;al (<xref ref-type="bibr" rid="B6">Bing zhi et&#xa0;al., 2010</xref>). in their study on hyperspectral estimation models of total nitrogen content in apple tree leaf leaves, which reflects the growth period adaptation law of fruit tree nutrition spectral diagnosis.</p>
</sec>
<sec id="s4_2">
<label>4.2</label>
<title>The synergistic mechanism of spectral preprocessing</title>
<p>Single preprocessing (MSC, SG, FD, SD) optimizes spectra from perspectives of physical interference elimination, noise suppression, dynamic enhancement, and fine feature extraction, yet exhibits functional limitations (e.g., FD tends to amplify noise, Second derivative is sensitive to noise) (<xref ref-type="bibr" rid="B5">Bao et&#xa0;al., 2024</xref>). Combined preprocessing (e.g., MS+FD, SG+SD achieves &#x201c;multi-functional synergy&#x201d;: MSC+FD first eliminates scattering interference and then amplifies chemical absorption dynamic change, enhancing the Spectral response of highly discrete LTP during the fruit-setting period; SG+SD first reduces noise to smooth the curve and then extracts fine structures of absorption peak, adapting to weak signals of low concentrations in the Maturity period. Correlation analysis shows that combined strategies can increase the r value of typical peaks and valleys by 0.05&#x2013;0.15, demonstrating that the &#x201c;synergistic effect&#x201d; of preprocessing is key to extracting Phosphorus-association information, providing an effective approach for spectral refinement of complex samples. This aligns with the consensus in the Chemometrics field that &#x201c;Combined preprocessing enhances Model performance&#x201d;, clarifying the refinement direction of spectral pre-processing in orchard Phosphorus diagnosis (<xref ref-type="bibr" rid="B28">Li et&#xa0;al., 2024</xref>).</p>
</sec>
<sec id="s4_3">
<label>4.3</label>
<title>Optimization of model hyper parameter and its impact on model performance</title>
<p>The performance of a Machine learning model is significantly influenced by the selection of hyper parameter, and the refinement of these hyper parameter directly affects the model&#x2019;s Generalization ability and Prediction accuracy (<xref ref-type="bibr" rid="B36">Schratz et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B56">Yang and Shami, 2020</xref>). Most existing studies have directly used default parameters to construct Spectroscopy estimation models without in-depth algorithmic refinement, which limits the performance improvement of the models. To address this limitation, this study systematically conducted research on hyper parameter refinement, employing grid search and cross validation methods to finely tune the parameters of Random forest (RF), Support vector machine (SVR), and BP neural network, significantly enhancing the stability and Prediction accuracy of the models.</p>
<p>For the RF model, experiments found that when n_estimators is 500 and min_samples_leaf is 5, the model achieves an optimal balance between training and prediction, effectively avoiding the phenomena of overfitting and poor fitting. In the Support Vector Regression model, the Penalty coefficient (C) C = 5 and kernel parameter &#x3b3;=0.1 were determined, and the Radial Basis Function (RBF) was selected. The Model performance outperformed other configurations such as linear kernel and polynomial kernel.</p>
<p>In the BP neural network, by comparing various training functions, the trainlm function was ultimately identified as the most suitable for the research task, achieving an ideal Fitting effect while ensuring convergence speed.</p>
<p>The above optimization results indicate that conducting parameter optimization for different algorithm systems can effectively exploit the model&#x2019;s potential and avoid the performance shortcomings caused by directly using default parameters. Through meticulous parameter tuning, this study provides a reliable configuration foundation for building a high-accuracy LTP content prediction model, and also offers a reference for algorithm optimization in similar Spectral modeling research.To ensure optimal Model performance, this study employed grid search and cross validation to refine the hyper parameters of RF, Support Vector Regression, and BP. Neither overfitting nor poor fitting phenomena occurred, further confirming the excellent performance of the model under this parameter combination. Therefore, the optimal parameters for this model are a number of decision trees of 500 and a min_samples_leaf of 5. Consequently, it is concluded that the model performs best when C = 5 and &#x3b3;=0.1. To further establish the model, the results indicated that the performance of the Radial Basis Function (RBF) is superior to the other three functions. Under the results, it was found that the trainlm training function is more suitable for this model.</p>
</sec>
<sec id="s4_4">
<label>4.4</label>
<title>Comparison with advanced baseline models</title>
<p>In previous studies, Partial least squares regression (PLSR), Light Gradient Boosting Machine (LightGBM) (LightGBM), and One-dimensional convolutional neural network (1D-CNN) have been widely used in the field of Spectroscopy. For example, researchers such as those from AgResearch employed ryegrass as experimental material and constructed a Spectroscopy prediction model using PLSR to evaluate the composition of ryegrass plants. Their results demonstrated strong predictive performance for total polysaccharide (R&#xb2; = 0.58), High molecular weight sugars (R&#xb2; = 0.63), ash (R&#xb2; = 0.50), and nitrogen content (R&#xb2; = 0.70) (<xref ref-type="bibr" rid="B39">Shorten et&#xa0;al., 2019</xref>). In another study, Jun Yan et&#xa0;al. used maize to develop a LightGBM-based prediction model for genomic selection prediction of maize lines. The model achieved an Area Under the Curve (AUC) of 0.793, indicating excellent performance in classification tasks involving large sample sizes (<xref ref-type="bibr" rid="B54">Yan et&#xa0;al., 2021</xref>).Guo, C. et&#xa0;al. constructed a cotton Fv/Fm prediction model based on 1D-CNN for drought tolerance assessment, using cotton as the experimental material. The predicted value showed a strong Correlation with the measured value (R&#xb2; &#x2265; 0.641). The results demonstrate that 1D-CNN offers high accuracy and stability in processing Large-scale data (<xref ref-type="bibr" rid="B20">Guo et&#xa0;al., 2022</xref>).</p>
<p>Although these models have shown excellent performance in studies by various researchers, the stability of PLSR under specific conditions across Different growth stages requires further enhancement. Both LightGBM and 1D-CNN are prone to high squared errors or significant bias when training samples are limited, increasing the risk of poor fitting. In comparison, this study identified more adaptive optimal Modeling strategies for Spectral features at Different growth stages through rigorous screening: the fruit-setting stage employs the FD+CARS-BP model, the expansion stage uses the SG+FD-CARS-BP model, and the Maturity period favors the SG+SD-CARS-BP model. The results show that these models exhibit superior predictive stability and adaptability across Different growth stages, enabling them to better handle the challenges of Modeling with small sample sizes, thus improving the accuracy of component prediction during specific growth stages.</p>
</sec>
<sec id="s4_5">
<label>4.5</label>
<title>Model generalization ability and cross-period challenges</title>
<p>Cross-period model comparisons revealed that the special models exhibited an R2 value 0.05&#x2013;0.16 higher than that of the general model during this growth period, while the Root mean square error (RMSE) was 0.0029&#x2013;0.0079 lower. Due to its inability to adapt to the &#x201c;dynamic Spectroscopy fingerprint&#x201d; of LTP across Different growth stages (such as the high discrete peak during the Fruit-setting period and the weak signal peak during the Maturity period), when the Fruit-setting period FD+CARS-BP model was extended to the Fruit expansion period, the R2 of the validation set decreased from 0.88 to 0.78, reflecting the specificity of the &#x201c;Spectroscopy&#x2013;phosphorus content&#x201d; relationship across growth periods. In practical applications, it is necessary to switch models based on the growth period or explore intertemporal transfer learning strategies (such as fine-tuning parameters of a pre-trained model) to balance model accuracy and convenience. This provides practical references for the field application and Roll out of orchard Spectroscopy models and also clarifies the direction for future model refinement&#x2014;enhancing the model&#x2019;s adaptability to differences between growth periods.</p>
</sec>
<sec id="s4_6">
<label>4.6</label>
<title>Limitations of the study and future directions</title>
<p>This study focuses on near-infrared spectroscopy (4000&#x2013;10000 cm<sup>-</sup>&#xb9;), with insufficient exploration of phosphorus characteristic peak (such as P&#x2013;O bond stretching vibration, ~1000&#x2013;1300 cm<sup>-</sup>&#xb9;). Future work could integrate Mid-infrared spectroscopy to expand features&#x2019; dimension, while simultaneously refining preprocessing and model parameter. Moreover, model training relies on laboratory Spectral data, without fully accounting for interference from field environments (e.g., light, temperature) on spectra. It is necessary to develop field spectral correction models and incorporate dynamic parameter adjustments to enhance technical practicality, thereby promoting the transition of Spectral diagnosis technology from the laboratory to practical application and improving the technical system for precision nutrient management in fruit trees.</p>
<p>In summary, this study clarifies the &#x201c;Growth period specificity-Pretreatment synergy-model adaptation&#x201d; technical framework for predicting Korla fragrant pear LTP, demonstrating that Stage-based modeling combined with Combined preprocessing can significantly improve prediction accuracy, providing a scientific paradigm for precision nutrient management in fruit trees. Subsequent efforts need to strengthen the integration of multiple Spectroscopy and field validation to further promote the application of Spectroscopy technology in orchard production.</p>
</sec>
</sec>
<sec id="s5" sec-type="conclusion">
<label>5</label>
<title>Conclusion</title>
<p>This study systematically analyzed the Leaf total phosphorus (LTP) content of Korla Fragrant Pear using Near-Infrared Spectroscopy, established a prediction model based on Growth period characteristics, and significantly improved detection accuracy and model applicability. The main conclusions include:</p>
<p>The LTP content of Korla Fragrant Pear leaves showed significant differences across various growing stages. The content was lowest during the Fruit-setting period, with a left-skewed distribution ranging from 0.02% to 0.25%; it stabilized during the Fruit expansion period, with a median of approximately 0.15%; and peaked during the Maturity period, exhibiting a right-skewed distribution with a maximum value of 0.45%. spectral analysis revealed that Spectral features in the 4000&#x2013;5500 cm<sup>-</sup>&#xb9; and 5500&#x2013;7500 cm<sup>-</sup>&#xb9; ranges were closely correlated with phosphorus content, providing a basis for developing the prediction model. This study constructed an LTP prediction model adapted to Different growth stages. The optimal model for the fruit-setting period was FD+CARS-BP, with the Coefficient of determination (R&#xb2;) for the training sets and validation set being 0.89 and 0.88, respectively; the optimal model for the fruit expansion period was SG+FD-CARS-BP, with the Coefficient of determination (R&#xb2;) for the training sets and validation set being 0.86 and 0.83, respectively; the optimal model for the Maturity period was SG+SD-CARS-BP, with the Coefficient of determination (R&#xb2;) for the training sets and validation set being 0.86 and 0.85, respectively. The predictive performance of all stage-specific models was significantly better than that of the intertemporal general model, with the Coefficient of determination (R&#xb2;) increasing by 0.05&#x2013;0.16 and the Root mean square error (RMSE) decreasing by 0.0029&#x2013;0.0079. This has practical implications for precision fertilization management in orchards and provides a basis for subsequent research to further enrich the trait system by combining Mid-infrared spectroscopy technology and to develop calibration models for real field environments, thereby enhancing the practicality and roll-out value of the method.</p>
<p>
<italic>commercial or financial relationships that could be construed as a potential conflict of interest</italic>.</p>
</sec>
</body>
<back>
<sec id="s6" sec-type="data-availability">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/supplementary material. Further inquiries can be directed to the corresponding author.</p>
</sec>
<sec id="s7" sec-type="author-contributions">
<title>Author contributions</title>
<p>MY: Writing &#x2013; original draft. JZ: Writing &#x2013; original draft. YL: Writing &#x2013; original draft. WF: Writing &#x2013; original draft. LW: Writing &#x2013; original draft. HW: Writing &#x2013; original draft. JB: Writing &#x2013; review &amp; editing.</p>
</sec>
<sec id="s8" sec-type="funding-information">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research and/or publication of this article. This project is supported by the Strong Youth Science and Technology Talent Project of the Corps &#x201c;Research on nitrogen regulation of pear calyx desiccation/retention fertilization strategy&#x201d; (project number: 2022CB001-11).</p>
</sec>
<sec id="s9" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s10" sec-type="ai-statement">
<title>Generative AI statement</title>
<p>The author(s) declare that no Generative AI was used in the creation of this manuscript.</p>
<p>Any alternative text (alt text) provided alongside figures in this article has been generated by Frontiers with the support of artificial intelligence and reasonable efforts have been made to ensure accuracy, including review by the authors wherever possible. If&#xa0;you identify any issues, please contact us.</p>
</sec>
<sec id="s11" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors&#xa0;and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="s12" sec-type="supplementary-material">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fpls.2025.1666460/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fpls.2025.1666460/full#supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Table1.docx" id="SM1" mimetype="application/vnd.openxmlformats-officedocument.wordprocessingml.document"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ahmadi</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Emami</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Daccache</surname> <given-names>A.</given-names>
</name>
<name>
<surname>He</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Soil properties prediction for precision agriculture using visible and near-infrared spectroscopy: A systematic review and meta-analysis</article-title>. <source>Agronomy.</source> <volume>11</volume>, <elocation-id>433</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/agronomy11030433</pub-id>
</citation></ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ahmad Yasmin</surname> <given-names>N. S.</given-names>
</name>
<name>
<surname>Abdul Wahab</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Ismail</surname> <given-names>F. S.</given-names>
</name>
<name>
<surname>MaJ</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Halim</surname> <given-names>M. H. A.</given-names>
</name>
<name>
<surname>Anuar</surname> <given-names>A. N.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Support vector regression modelling of an aerobic granular sludge in sequential batch reactor</article-title>. <source>Membranes.</source> <volume>11</volume>, <elocation-id>554</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/membranes11080554</pub-id>, PMID: <pub-id pub-id-type="pmid">34436317</pub-id></citation></ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>An</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Yan</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Lu</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>X.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Effect of spectral pretreatment on qualitative identification of adulterated bovine colostrum by near-infrared spectroscopy</article-title>. <source>Infrared Phys. Technol.</source> <volume>118</volume>, <elocation-id>103869</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.infrared.2021.103869</pub-id>
</citation></ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Arias</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Zambrano</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Broce</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Medina</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Pacheco</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Nunez</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Hyperspectral imaging for rice cultivation: Applications, methods and challenges</article-title>. <source>AIMS Agric. Food</source> <volume>6</volume>, <fpage>273</fpage>&#x2013;<lpage>307</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3934/agrfood.2021018</pub-id>
</citation></ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bao</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Tang</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Zhi</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Determination of leaf nitrogen content in apple and jujube by near-infrared spectroscopy</article-title>. <source>Sci. Rep.</source> <volume>14</volume>, <fpage>20884</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41598-024-71590-1</pub-id>, PMID: <pub-id pub-id-type="pmid">39242639</pub-id></citation></ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bing zhi</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Ming xia</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Xuan</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Lin sen</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Hai yan</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Hyperspectral estimation models for nitrogen contents of apple leaves</article-title>. <source>J. Remote Sensing.</source> <volume>14</volume>, <fpage>761</fpage>&#x2013;<lpage>773</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.11834/jrs.20100411</pub-id>
</citation></ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cao</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Bao</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Wei</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Duan</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Du</surname> <given-names>Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>Fast performance modeling across different database versions using partitioned co-kriging</article-title>. <source>Appl. Sci.</source> <volume>11</volume>, <elocation-id>9669</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/app11209669</pub-id>
</citation></ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Capitaine</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Genuer</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Thi&#xe9;baut</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Random forests for high-dimensional longitudinal data</article-title>. <source>Stat. Methods Med. Res.</source> <volume>30</volume>, <fpage>166</fpage>&#x2013;<lpage>184</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1177/0962280220946080</pub-id>, PMID: <pub-id pub-id-type="pmid">32772626</pub-id></citation></ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Jin</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Physiological and morphological responses of hydroponically grown pear rootstock under phosphorus treatment [Original research</article-title>. <source>Front. Plant Sci.</source> <volume>12</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fpls.2021.696045</pub-id>, PMID: <pub-id pub-id-type="pmid">34858445</pub-id></citation></ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chu</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Wen</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Nan</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Du</surname> <given-names>C.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Possible alternatives: identifying and quantifying adulteration in buffalo, goat, and camel milk using mid-infrared spectroscopy combined with modern statistical machine learning methods</article-title>. <source>Foods.</source> <volume>12</volume>, <elocation-id>3856</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/foods12203856</pub-id>, PMID: <pub-id pub-id-type="pmid">37893749</pub-id></citation></ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dabiri</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Farhangi</surname> <given-names>V.</given-names>
</name>
<name>
<surname>Moradi</surname> <given-names>M. J.</given-names>
</name>
<name>
<surname>Zadehmohamad</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Karakouzian</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Applications of decision tree and random forest as tree-based machine learning techniques for analyzing the ultimate strain of spliced and non-spliced reinforcement bars</article-title>. <source>Appl. Sci.</source> <volume>12</volume>, <elocation-id>4851</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/app12104851</pub-id>
</citation></ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dayton</surname> <given-names>E. A.</given-names>
</name>
<name>
<surname>Whitacre</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Holloman</surname> <given-names>C.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Comparison of three persulfate digestion methods for total phosphorus analysis and estimation of suspended sediments</article-title>. <source>Appl. Geochemistry</source> <volume>78</volume>, <fpage>357</fpage>&#x2013;<lpage>362</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.apgeochem.2017.01.011</pub-id>
</citation></ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fonseca-Garc&#xed;a</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Solis-Miranda</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Pacheco</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Quinto</surname> <given-names>C.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Non-specific lipid transfer proteins in legumes and their participation during root-nodule symbiosis [Original research</article-title>. <source>Front. Agron.</source>, <fpage>3</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fagro.2021.660100</pub-id>
</citation></ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gao</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Xiao</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Bao</surname> <given-names>T.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Detection and identification of potato-typical diseases based on multidimensional fusion atrous-CNN and hyperspectral data</article-title>. <source>Appl. Sci.</source> <volume>13</volume>, <elocation-id>5023</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/app13085023</pub-id>
</citation></ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gao</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Cheng</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>G.</given-names>
</name>
<etal/>
</person-group>. (<year>2016</year>). <article-title>Improve the prediction accuracy of apple tree canopy nitrogen content through multiple scattering correction using spectroscopy</article-title>. <source>Agric. Sci.</source> <volume>010)</volume>, <fpage>007</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.4236/as.2016.710061</pub-id>
</citation></ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gautam</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Vanga</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Ariese</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Umapathy</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Review of multidimensional data processing approaches for Raman and infrared spectroscopy</article-title>. <source>EPJ Techniques Instrumentation</source> <volume>2</volume>, <elocation-id>8</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1140/epjti/s40485-015-0018-6</pub-id>
</citation></ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guo</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Feng</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>W.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Estimation of potato canopy leaf water content in various growth stages using UAV hyperspectral remote sensing and machine learning [Original Research</article-title>. <source>Front. Plant Sci.</source> <volume>15</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fpls.2024.1458589</pub-id>, PMID: <pub-id pub-id-type="pmid">39610888</pub-id></citation></ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guo</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Hao</surname> <given-names>P.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Using a random forest model to predict the location of potential damage on asphalt pavement</article-title>. <source>Appl. Sci.</source> <volume>11</volume>, <elocation-id>10396</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/app112110396</pub-id>
</citation></ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guo</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Qin</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>B.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Mechanical fault diagnosis of a DC motor utilizing united variational mode decomposition, sampEn, and random forest-SPRINT algorithm classifiers</article-title>. <source>Entropy.</source> <volume>21</volume>, <elocation-id>470</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/e21050470</pub-id>, PMID: <pub-id pub-id-type="pmid">33267184</pub-id></citation></ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guo</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Sun</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Predicting Fv/Fm and evaluating cotton drought tolerance using hyperspectral and 1D-CNN [Original Research</article-title>. <source>Front. Plant Sci.</source> <volume>13</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fpls.2022.1007150</pub-id>, PMID: <pub-id pub-id-type="pmid">36330250</pub-id></citation></ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gupta</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Kulkarni</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Mukherjee</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Accurate prediction of B-form/A-form DNA conformation propensity from primary sequence: A machine learning and free energy handshake</article-title>. <source>Patterns</source> <volume>2</volume>, <elocation-id>100329</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.patter.2021.100329</pub-id>, PMID: <pub-id pub-id-type="pmid">34553171</pub-id></citation></ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Han</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Jiang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Sun</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2025</year>). <article-title>Non-destructive identification of commercial jerky types based on multi-band hyperspectral imaging with machine learning</article-title>. <source>Food Chemistry: X</source> <volume>26</volume>, <elocation-id>102293</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.fochx.2025.102293</pub-id>, PMID: <pub-id pub-id-type="pmid">40071142</pub-id></citation></ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>He</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>S.-B.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Y.-Z.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>An integrated chemical characterization based on FT-NIR, and GC&#x2013;MS for the comparative metabolite profiling of 3 species of the genus Amomum</article-title>. <source>Analytica Chimica Acta</source> <volume>1280</volume>, <elocation-id>341869</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.aca.2023.341869</pub-id>, PMID: <pub-id pub-id-type="pmid">37858569</pub-id></citation></ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Hu</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Cai</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>D.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Short term electrical load forecasting using mutual information based feature selection with generalized minimum-redundancy and maximum-relevance criteria</article-title>. <source>Entropy.</source> <volume>18</volume>, <elocation-id>330</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/e18090330</pub-id>
</citation></ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jeong</surname> <given-names>J. H.</given-names>
</name>
<name>
<surname>Resop</surname> <given-names>J. P.</given-names>
</name>
<name>
<surname>Mueller</surname> <given-names>N. D.</given-names>
</name>
<name>
<surname>Fleisher</surname> <given-names>D. H.</given-names>
</name>
<name>
<surname>Yun</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Butler</surname> <given-names>E. E.</given-names>
</name>
<etal/>
</person-group>. (<year>2016</year>). <article-title>Random forests for global and regional crop yield predictions</article-title>. <source>PloS One</source> <volume>11</volume>, <fpage>e0156571</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1371/journal.pone.0156571</pub-id>, PMID: <pub-id pub-id-type="pmid">27257967</pub-id></citation></ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Jin</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Dogani</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Gu</surname> <given-names>X.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Enhancing lightGBM for industrial fault warning: an innovative hybrid algorithm</article-title>. <source>Processes.</source> <volume>12</volume>, <elocation-id>221</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/pr12010221</pub-id>
</citation></ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Dai</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Qiu</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Pang</surname> <given-names>X.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Application of PLSR in correlating sensory and chemical properties of middle flue-cured tobacco leaves with honey-sweet and burnt flavour</article-title>. <source>Heliyon</source> <volume>10</volume>, <elocation-id>e29547</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.heliyon.2024.e29547</pub-id>, PMID: <pub-id pub-id-type="pmid">38655300</pub-id></citation></ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Peng</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Yin</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>B.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Optimization of online soluble solids content detection models for apple whole fruit with different mode spectra combined with spectral correction and model fusion</article-title>. <source>Foods.</source> <volume>13</volume>, <elocation-id>1037</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/foods13071037</pub-id>, PMID: <pub-id pub-id-type="pmid">38611343</pub-id></citation></ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Tan</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>J.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>A particle swarm optimization improved BP neural network intelligent model for electrocardiogram classification</article-title>. <source>BMC Med. Inf. Decision Making</source> <volume>21</volume>, <fpage>99</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s12911-021-01453-6</pub-id>, PMID: <pub-id pub-id-type="pmid">34330266</pub-id></citation></ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li Ting</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Hong Xing</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Research on determination of total phosphorus content in water treatment agent by molybdenum-antimony-ascorbic acid method</article-title>. <source>ANHUI Chem. INDUSTRY.</source> <volume>48</volume>, <fpage>114</fpage>&#x2013;<lpage>119</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3969/j.issn.1008-553X.2022.05.028</pub-id>
</citation></ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Dang</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Lin</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Applications of savitzky-golay filter for seismic random noise reduction</article-title>. <source>Acta Geophysica</source> <volume>64</volume>, <fpage>101</fpage>&#x2013;<lpage>124</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1515/acgeo-2015-0062</pub-id>
</citation></ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>L.-L.</given-names>
</name>
<name>
<surname>Lu</surname> <given-names>Y.-P.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Gu</surname> <given-names>X.-Y.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>L.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Deep_KsuccSite: A novel deep learning method for the identification of lysine succinylation sites [Original Research</article-title>. <source>Front. Genet.</source> <volume>13</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fgene.2022.1007618</pub-id>, PMID: <pub-id pub-id-type="pmid">36246655</pub-id></citation></ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Murguzur</surname> <given-names>F. J. A.</given-names>
</name>
<name>
<surname>Bison</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Smis</surname> <given-names>A.</given-names>
</name>
<name>
<surname>B&#xf6;hner</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Struyf</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Meire</surname> <given-names>P.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>Towards a global arctic-alpine model for Near-infrared reflectance spectroscopy (NIRS) predictions of foliar nitrogen, phosphorus and carbon content</article-title>. <source>Sci. Rep.</source> <volume>9</volume>, <fpage>8259</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41598-019-44558-9</pub-id>, PMID: <pub-id pub-id-type="pmid">31164672</pub-id></citation></ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qi</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Hu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Feng</surname> <given-names>W.</given-names>
</name>
<etal/>
</person-group>. (<year>2025</year>). <article-title>Classification of different gluten wheat varieties based on hyperspectral preprocessing, feature screening, and machine learning</article-title>. <source>Food Chemistry: X</source> <volume>26</volume>, <elocation-id>102329</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.fochx.2025.102329</pub-id>, PMID: <pub-id pub-id-type="pmid">40123867</pub-id></citation></ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rhodes</surname> <given-names>J. S.</given-names>
</name>
<name>
<surname>Cutler</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Moon</surname> <given-names>K. R.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Geometry- and accuracy-preserving random forest proximities</article-title>. <source>IEEE Trans. Pattern Anal. Mach. Intell.</source> <volume>45</volume>, <fpage>10947</fpage>&#x2013;<lpage>10959</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/tpami.2023.3263774</pub-id>, PMID: <pub-id pub-id-type="pmid">37015125</pub-id></citation></ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schratz</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Muenchow</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Iturritxa</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Richter</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Brenning</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Hyperparameter tuning and performance assessment of statistical and machine-learning algorithms using spatial data</article-title>. <source>Ecol. Model.</source> <volume>406</volume>, <fpage>109</fpage>&#x2013;<lpage>120</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.ecolmodel.2019.06.002</pub-id>
</citation></ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Scornet</surname> <given-names>E.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>On the asymptotics of random forests</article-title>. <source>J. Multivariate Anal.</source> <volume>146</volume>, <fpage>72</fpage>&#x2013;<lpage>83</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.jmva.2015.06.009</pub-id>
</citation></ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shah</surname> <given-names>J. A.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Ullah</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Duan</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Shen</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Liao</surname> <given-names>Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2024</year>). <article-title>Linkages among leaf nutrient concentration, resorption efficiency, litter decomposition and their stoichiometry to canopy nitrogen addition and understory removal in subtropical plantation</article-title>. <source>Ecol. Processes</source> <volume>13</volume>, <fpage>27</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s13717-024-00507-7</pub-id>
</citation></ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shorten</surname> <given-names>P. R.</given-names>
</name>
<name>
<surname>Leath</surname> <given-names>S. R.</given-names>
</name>
<name>
<surname>Schmidt</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Ghamkhar</surname> <given-names>K.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Predicting the quality of ryegrass using hyperspectral imaging</article-title>. <source>Plant Methods</source> <volume>15</volume>, <fpage>63</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s13007-019-0448-2</pub-id>, PMID: <pub-id pub-id-type="pmid">31182971</pub-id></citation></ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Siedliska</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Baranowski</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Pastuszka-Wo&#x17a;niak</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Zubik</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Krzyszczak</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Identification of plant leaf phosphorus content at different growth stages based on hyperspectral reflectance</article-title>. <source>BMC Plant Biol.</source> <volume>21</volume>, <elocation-id>28</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s12870-020-02807-4</pub-id>, PMID: <pub-id pub-id-type="pmid">33413120</pub-id></citation></ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Song</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Lin</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Wei</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Zeng</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Mu</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>X.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Trichoderma viride improves phosphorus uptake and the growth of Chloris virgata under phosphorus-deficient conditions [Original Research</article-title>. <source>Front. Microbiol.</source> <volume>15</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fmicb.2024.1425034</pub-id>, PMID: <pub-id pub-id-type="pmid">39027109</pub-id></citation></ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sonobe</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Hirono</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Applying variable selection methods and preprocessing techniques to hyperspectral reflectance data to estimate tea cultivar chlorophyll content</article-title>. <source>Remote Sensing.</source> <volume>15</volume>, <elocation-id>19</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/rs15010019</pub-id>
</citation></ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sun</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Feng</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Xiao</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Ding</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Estimation of water content in corn leaves using hyperspectral data based on fractional order Savitzky-Golay derivation coupled with wavelength selection</article-title>. <source>Comput. Electron. Agric.</source> <volume>182</volume>, <elocation-id>105989</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compag.2021.105989</pub-id>
</citation></ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tan</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>X.</given-names>
</name>
</person-group> (<year>2025</year>). <article-title>A hyperspectral feature selection method for soil organic matter estimation based on an improved weighted marine predators algorithm</article-title>. <source>IEEE Trans. Geosci. Remote Sensing.</source> <volume>63</volume>, <fpage>1</fpage>&#x2013;<lpage>11</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/TGRS.2024.3422475</pub-id>
</citation></ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tian</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>He</surname> <given-names>Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2024</year>). <article-title>New 3-D fluorescence spectral indices for multiple pigment inversions of plant leaves via 3-D fluorescence spectra</article-title>. <source>Remote Sensing.</source> <volume>16</volume>, <elocation-id>1885</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/rs16111885</pub-id>
</citation></ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Virnodkar</surname> <given-names>S. S.</given-names>
</name>
<name>
<surname>Pachghare</surname> <given-names>V. K.</given-names>
</name>
<name>
<surname>Patil</surname> <given-names>V. C.</given-names>
</name>
<name>
<surname>Jha</surname> <given-names>S. K.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Remote sensing and machine learning for crop water stress determination in various crops: a critical review</article-title>. <source>Precis. Agric.</source> <volume>21</volume>, <fpage>1121</fpage>&#x2013;<lpage>1155</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s11119-020-09711-9</pub-id>
</citation></ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Feng</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Gui</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Teng</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Ye</surname> <given-names>F.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Digital prediction of the purchase price of fresh tea leaves of enshi yulu based on near-infrared spectroscopy combined with multivariate analysis</article-title>. <source>Foods.</source> <volume>12</volume>, <elocation-id>3592</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/foods12193592</pub-id>, PMID: <pub-id pub-id-type="pmid">37835242</pub-id></citation></ref>
<ref id="B48">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>He</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Gong</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Z.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Optimization of a water-saving and fertilizer-saving model for enhancing xinjiang korla fragrant pear yield, quality, and net profits under water and fertilizer coupling</article-title>. <source>Sustainability.</source> <volume>14</volume>, <elocation-id>8495</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/su14148495</pub-id>
</citation></ref>
<ref id="B49">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Hou</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Sobhy</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Prediction and optimization of tower mill grinding power consumption based on GA-BP neural network</article-title>. <source>Physicochem Probl Miner Process.</source> <volume>59</volume>, <elocation-id>172096</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.37190/ppmp/172096</pub-id>
</citation></ref>
<ref id="B50">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xiao</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Tang</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Feng</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>L.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Visible and near-infrared spectroscopy and deep learning application for the qualitative and quantitative investigation of nitrogen status in cotton leaves [Original Research</article-title>. <source>Front. Plant Sci.</source> <volume>13</volume>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fpls.2022.1080745</pub-id>, PMID: <pub-id pub-id-type="pmid">36643292</pub-id></citation></ref>
<ref id="B51">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xie</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>F.-Y.</given-names>
</name>
<name>
<surname>Fan</surname> <given-names>X.-J.</given-names>
</name>
<name>
<surname>Hu</surname> <given-names>S.-J.</given-names>
</name>
<name>
<surname>Xiao</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>J.-F.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Components analysis of biochar based on near infrared spectroscopy technology</article-title>. <source>Chin. J. Analytical Chem.</source> <volume>46</volume>, <fpage>609</fpage>&#x2013;<lpage>615</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/S1872-2040(17)61081-8</pub-id>
</citation></ref>
<ref id="B52">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xing</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Du</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Shen</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A method combining FTIR-ATR and Raman spectroscopy to determine soil organic matter: Improvement of prediction accuracy using competitive adaptive reweighted sampling (CARS)</article-title>. <source>Comput. Electron. Agric.</source> <volume>191</volume>, <elocation-id>106549</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compag.2021.106549</pub-id>
</citation></ref>
<ref id="B53">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>Z.-F.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.-Y.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Deng</surname> <given-names>G.-P.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>N.</given-names>
</name>
<etal/>
</person-group>. (<year>2016</year>). <article-title>Leaf decomposition and nutrient release of dominant species in the forest and lake in the Jiuzhaigou National Nature Reserve, China</article-title>. <source>Chin. J. Plant Ecol.</source> <volume>40</volume>, <fpage>883</fpage>&#x2013;<lpage>892</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.17521/cjpe.2016.0040</pub-id>
</citation></ref>
<ref id="B54">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yan</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Cheng</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Jiang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Xiao</surname> <given-names>Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2021</year>). <article-title>LightGBM: accelerated genomically designed crop breeding through ensemble learning</article-title>. <source>Genome Biol.</source> <volume>22</volume>, <fpage>271</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s13059-021-02492-y</pub-id>, PMID: <pub-id pub-id-type="pmid">34544450</pub-id></citation></ref>
<ref id="B55">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yang</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Shi</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Sun</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Studies on fault diagnosis of dissolved oxygen sensor based on GA-SVM</article-title>. <source>Math. Biosci. Engineering.</source> <volume>18</volume>, <fpage>386</fpage>&#x2013;<lpage>399</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3934/mbe.2021021</pub-id>, PMID: <pub-id pub-id-type="pmid">33525098</pub-id></citation></ref>
<ref id="B56">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yang</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Shami</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>On hyperparameter optimization of machine learning algorithms: Theory and practice</article-title>. <source>Neurocomputing</source> <volume>415</volume>, <fpage>295</fpage>&#x2013;<lpage>316</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.neucom.2020.07.061</pub-id>
</citation></ref>
<ref id="B57">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yu</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Bai</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Bao</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Tang</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Zheng</surname> <given-names>Q.</given-names>
</name>
<etal/>
</person-group>. (<year>2024</year>). <article-title>The prediction&#xa0;model of total nitrogen content in leaves of korla fragrant pear was established based on near infrared spectroscopy</article-title>. <source>Agronomy.</source> <volume>14</volume>, <elocation-id>1284</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/agronomy14061284</pub-id>
</citation></ref>
<ref id="B58">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Chang</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Hyperspectral estimation of chlorophyll content in apple tree leaf based on feature band selection and the catBoost model</article-title>. <source>Agronomy.</source> <volume>13</volume>, <elocation-id>2075</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/agronomy13082075</pub-id>
</citation></ref>
<ref id="B59">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Gao</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Cen</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Lu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>He</surname> <given-names>Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>Automated spectral feature extraction from hyperspectral images to differentiate weedy rice and barnyard grass from a rice crop</article-title>. <source>Comput. Electron. Agric.</source> <volume>159</volume>, <fpage>42</fpage>&#x2013;<lpage>49</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compag.2019.02.018</pub-id>
</citation></ref>
<ref id="B60">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Shen</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>S.</given-names>
</name>
<name>
<surname>He</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Application of near-infrared spectroscopy for the nondestructive analysis of wheat flour: A review</article-title>. <source>Curr. Res. Food Sci.</source> <volume>5</volume>, <fpage>1305</fpage>&#x2013;<lpage>1312</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.crfs.2022.08.006</pub-id>, PMID: <pub-id pub-id-type="pmid">36065198</pub-id></citation></ref>
<ref id="B61">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zheng</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Data fusion of FT-NIR and ATR-FTIR spectra for accurate authentication of geographical indications for Gastrodia elata Blume</article-title>. <source>Food Bioscience</source> <volume>56</volume>, <elocation-id>103308</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.fbio.2023.103308</pub-id>
</citation></ref>
<ref id="B62">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zheng</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Effect of drying temperature on composition of edible mushrooms: Characterization and assessment via HS-GC-MS and IR spectral based volatile profiling and chemometrics</article-title>. <source>Curr. Res. Food Sci.</source> <volume>9</volume>, <elocation-id>100819</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.crfs.2024.100819</pub-id>, PMID: <pub-id pub-id-type="pmid">39234276</pub-id></citation></ref>
<ref id="B63">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Su</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>He</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Fang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Sun</surname> <given-names>X.</given-names>
</name>
<etal/>
</person-group>. (<year>2024</year>). <article-title>Aggregation and assessment of grape quality parameters with visible-near-infrared spectroscopy: Introducing a novel quantitative index</article-title>. <source>Postharvest Biol. Technol.</source> <volume>218</volume>, <elocation-id>113131</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.postharvbio.2024.113131</pub-id>
</citation></ref>
<ref id="B64">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zou</surname> <given-names>X.-H.</given-names>
</name>
<name>
<surname>Hu</surname> <given-names>Y.-N.</given-names>
</name>
<name>
<surname>Wei</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>S.-T.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>P.-F.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>X.-Q.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Correlation between endogenous hormone and the adaptability of Chinese fir with high phosphorus-use efficiency to low phosphorus stress</article-title>. <source>Chin. J. Plant Ecol.</source> <volume>43</volume>, <fpage>139</fpage>&#x2013;<lpage>151</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.17521/cjpe.2018.0201</pub-id>
</citation></ref>
</ref-list>
</back>
</article>