<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Plant Sci.</journal-id>
<journal-title>Frontiers in Plant Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Plant Sci.</abbrev-journal-title>
<issn pub-type="epub">1664-462X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpls.2022.841452</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Plant Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Soluble Solids Content Binary Classification of Miyagawa Satsuma in Chongming Island Based on Near Infrared Spectroscopy</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name><surname>Chen</surname> <given-names>Yuzhen</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x02020;</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Sun</surname> <given-names>Wanxia</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1806082/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Jiu</surname> <given-names>Songtao</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/850749/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Wang</surname> <given-names>Lei</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1403999/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Deng</surname> <given-names>Bohan</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Chen</surname> <given-names>Zili</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Jiang</surname> <given-names>Fei</given-names></name>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Hu</surname> <given-names>Menghan</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x0002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1170778/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Zhang</surname> <given-names>Caixi</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c002"><sup>&#x0002A;</sup></xref>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>School of Agriculture and Biology, Shanghai Jiao Tong University</institution>, <addr-line>Shanghai</addr-line>, <country>China</country></aff>
<aff id="aff2"><sup>2</sup><institution>Shanghai Key Laboratory of Multidimensional Information Processing, School of Communication and Electronic Engineering, East China Normal University</institution>, <addr-line>Shanghai</addr-line>, <country>China</country></aff>
<aff id="aff3"><sup>3</sup><institution>Shanghai Citrus Research Institute</institution>, <addr-line>Shanghai</addr-line>, <country>China</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: Chuanlei Zhang, Tianjin University of Science and Technology, China</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Ke Han, Harbin University of Commerce, China; Xing Wei, Purdue University, United States</p></fn>
<corresp id="c001">&#x0002A;Correspondence: Menghan Hu <email>mhhu&#x00040;ce.ecnu.edu.cn</email></corresp>
<corresp id="c002">Caixi Zhang <email>acaizh&#x00040;sjtu.edu.cn</email></corresp>
<fn fn-type="other" id="fn001"><p>This article was submitted to Sustainable and Intelligent Phytoprotection, a section of the journal Frontiers in Plant Science</p></fn>
<fn fn-type="equal" id="fn002"><p>&#x02020;These authors have contributed equally to this work</p></fn></author-notes>
<pub-date pub-type="epub">
<day>18</day>
<month>07</month>
<year>2022</year>
</pub-date>
<pub-date pub-type="collection">
<year>2022</year>
</pub-date>
<volume>13</volume>
<elocation-id>841452</elocation-id>
<history>
<date date-type="received">
<day>22</day>
<month>12</month>
<year>2021</year>
</date>
<date date-type="accepted">
<day>15</day>
<month>06</month>
<year>2022</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x000A9; 2022 Chen, Sun, Jiu, Wang, Deng, Chen, Jiang, Hu and Zhang.</copyright-statement>
<copyright-year>2022</copyright-year>
<copyright-holder>Chen, Sun, Jiu, Wang, Deng, Chen, Jiang, Hu and Zhang</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license></permissions>
<abstract>
<p>Citrus is one of the most important fruits in China. Miyagawa Satsuma, one kind of citrus, is a nutritious agricultural product with regional characteristics of Chongming Island. Near-infrared Spectroscopy (NIR) is a proper method for studying the quality of fruits, because it is low-cost, efficient, non-destructive, and repeatable. Therefore, the NIR technique is used to detect citrus&#x00027;s soluble solid content (SSC) in this study. After obtaining the original spectral data, the first 70% of them are divided into the training set and 30% into the test set. Then, the Random Frog algorithm is chosen to select characteristic wavelengths, which reduces the dimension of the data and the complexity of the model, and accordingly makes the generalization of the classification model better. After comparing the performance of various classifiers (AdaBoost, KNN, LS-SVM, and Bayes) under different characteristic wavelength numbers, the AdaBoost classifier outperforms using 275 characteristic wavelengths for modeling eventually. The accuracy, precision, recall, and <italic>F</italic><sub>1</sub>-score are 78.3%, 80.5%, 78.3%, and 0.780, respectively and the ROC (Receiver Operating Characteristic Curve, ROC curve) is close to the upper left corner, suggesting that the classification model is acceptable. The results demonstrate that it is feasible to use the NIR technique to estimate whether the citrus is sweet or not. Furthermore, it is beneficial for us to apply the obtained models for identifying the quality of citrus correctly. For fruit traders, the model helps them to determine the growth cycle of citrus more scientifically, improve the level of citrus cultivation and management and the final fruit quality, and thus increase the economic income of fruit traders.</p></abstract>
<kwd-group>
<kwd>near infrared spectroscopy</kwd>
<kwd>AdaBoost</kwd>
<kwd>random frog</kwd>
<kwd>citrus soluble solids content</kwd>
<kwd>machine learning</kwd>
</kwd-group>
<counts>
<fig-count count="5"/>
<table-count count="1"/>
<equation-count count="4"/>
<ref-count count="36"/>
<page-count count="10"/>
<word-count count="6631"/>
</counts>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<title>1. Introduction</title>
<p>Citrus fruits are among the most commonly grown and consumed fruits all over the world and meanwhile one of the most important fruits in China since they are very nutritious and can supplement vitamins, promote digestion and increase appetite (Zou et al., <xref ref-type="bibr" rid="B36">2016</xref>; Anticona et al., <xref ref-type="bibr" rid="B1">2020</xref>). The total output of <italic>Citrus reticulata Blanco</italic> is 21.2 million tons in China, accounting for 67% of the total citrus output. <italic>Citrus unshiu</italic> is one of the three main varieties of citrus reticulata Blanco in China (Nam et al., <xref ref-type="bibr" rid="B27">2019</xref>; Cheng et al., <xref ref-type="bibr" rid="B7">2020</xref>). This research uses Miyagawa Satsuma, a variety of <italic>citrus unshiu</italic>, from Chongming Island in Shanghai, as the research object. The citrus in Chongming Island not only grows in environmental conditions famous for fresh air, clean water, and rich soil but also ripens in cultivation technology of &#x0201C;green prevention and control and plastic film covering with grass and organic cultivation&#x0201D; (this cultivation concept originates from Shanghai Qianwei Citrus Co., Ltd.). As a result, it owns the advantages of both rich nutrition and the regional characteristics of Chongming Island.</p>
<p>Soluble solids content (SSC), is one of the most important internal quality attributes of most fruits. The SSC plays an important role in the fruit maturity process and partly influences the flavor of most fruits, thus determining the acceptance of rich nutrients and economic benefits in the fruit trade. The detection of citrus SSC is not only beneficial to customers but also significant for growers (Li et al., <xref ref-type="bibr" rid="B20">2016b</xref>; Fan et al., <xref ref-type="bibr" rid="B10">2019</xref>; Guo et al., <xref ref-type="bibr" rid="B13">2020</xref>). Therefore, in recent decades, the demand to develop non-destructive and rapid evaluation methods for citrus SSC has become more extensive and urgent. Electronic nose technology (Zhang et al., <xref ref-type="bibr" rid="B33">2008</xref>, <xref ref-type="bibr" rid="B34">2016</xref>), computer vision (Xia et al., <xref ref-type="bibr" rid="B30">2016</xref>; Bhargava and Bansal, <xref ref-type="bibr" rid="B5">2021</xref>), and hyperspectral imaging technology (Li et al., <xref ref-type="bibr" rid="B18">2016a</xref>, <xref ref-type="bibr" rid="B22">2018</xref>) are some common methods to measure the quality of fruits. However, electronic nose technology is restricted to limited enclosed space, which is inconvenient to carry out. Computer vision technology lacks spectral information. As for hyperspectral imaging technology, the obtained hypercube contains a lot of redundant information which leads to a high computation cost.</p>
<p>Fortunately, with the advantage of low testing cost, high efficiency, good reproducibility of test results, and non-destructive testing, NIR spectroscopy, between wavelength region range of 780&#x02013;2,526 nm, has been applied popularly in the analysis of different fruit or vegetable samples (Beghi et al., <xref ref-type="bibr" rid="B4">2017</xref>; Arendse et al., <xref ref-type="bibr" rid="B3">2018</xref>), such as apple (Xia et al., <xref ref-type="bibr" rid="B29">2020</xref>; Arefi et al., <xref ref-type="bibr" rid="B2">2021</xref>; Ma et al., <xref ref-type="bibr" rid="B24">2021</xref>; Li et al., <xref ref-type="bibr" rid="B21">2022</xref>), tomato (Huang et al., <xref ref-type="bibr" rid="B14">2021</xref>; Zhang et al., <xref ref-type="bibr" rid="B32">2021</xref>), persimmon (Wei et al., <xref ref-type="bibr" rid="B28">2020</xref>), pear (Cruz et al., <xref ref-type="bibr" rid="B8">2021</xref>), and banana (Cruz et al., <xref ref-type="bibr" rid="B8">2021</xref>). Xia et al. (<xref ref-type="bibr" rid="B29">2020</xref>) studied the effect of sample diameter differences on the online prediction of SSC of &#x0201C;Fuji&#x0201D; apples with the methods of visible and near-infrared spectroscopy and partial least square regression. It is justified that diffuse transmission spectra in 710&#x02013;980 nm and diameter correction method with calculated attenuation coefficient are the best. Wei et al. used NIR hyperspectral imaging within 900&#x02013;1,700 nm to model SSC and firmness determination of persimmon with partial least squares regression (Wei et al., <xref ref-type="bibr" rid="B28">2020</xref>). The final models obtained a coefficient of determination of 0.757, RMSEP of 1.404 Brix, and <inline-formula><mml:math id="M1"><mml:msubsup><mml:mrow><mml:mi>R</mml:mi></mml:mrow><mml:mrow><mml:mi>p</mml:mi></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:msubsup></mml:math></inline-formula> of 0.876, RMSEP of 0.395 for SSC and firmness detection, respectively. Pahlawan et al. developed the calibration model to predict the SSC of bananas using NIR spectroscopy in the range from 350 to 1,000 nm. It was conducted by various distances of fiber optic probes to bananas samples (Cruz et al., <xref ref-type="bibr" rid="B8">2021</xref>). To our best knowledge, there have been few similar studies on Miyagawa Satsuma. From these researches mentioned above, it can be easily seen that they focus on predicting the accurate number of the attribute focused on, such as SSC. Sometimes, we are more interested in knowing the sugar level rating rather than the specific value. There is no exact numerical index for distinguishing sweet and unsweetened, and therefore, we calculate the average value of SSC as the demarcation index for judging sweet or unsweet for the reason that SSC is an important index affecting the sweetness.</p>
<p>Reducing dimensions and seeking the most informative wavelengths are effective methods for processing data while selecting the most informative wavelengths of target information is an effective measure to simplify computation and improve the model performance (Li et al., <xref ref-type="bibr" rid="B19">2019</xref>; Zhou et al., <xref ref-type="bibr" rid="B35">2020</xref>). First, it has been shown that the inclusion of uninformative wavelengths while modeling affects the performance of predicting or classifying and model interpretability (Chang et al., <xref ref-type="bibr" rid="B6">2016</xref>). Second, the identification of wavelengths that contain information about the attribute the research focuses on, will reduce the computation time and cost, from a more practical point of view (Zhang et al., <xref ref-type="bibr" rid="B31">2019</xref>; Mamouei et al., <xref ref-type="bibr" rid="B25">2020</xref>). Li et al. (<xref ref-type="bibr" rid="B18">2016a</xref>) chose the carlo-uninformative variable elimination and successive projections algorithm to select the most effective variables from hyperspectral data when doing the research on measuring SSC in pear. The results indicated that the model built using 18 effective variables achieved the optimal performance for the prediction of SSC. Jun et al. (<xref ref-type="bibr" rid="B15">2018</xref>) used an iteratively retaining informative variables algorithm to obtain 10 characteristic wavelengths when processing samples in predicting the SSC of cherry tomatoes. The experimental results showed the IRIV&#x02212;CS&#x02212;SVR model for SSC prediction could reach accuracy with <inline-formula><mml:math id="M2"><mml:msubsup><mml:mrow><mml:mi>R</mml:mi></mml:mrow><mml:mrow><mml:mi>p</mml:mi></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:msubsup></mml:math></inline-formula> = 0.9718 and <inline-formula><mml:math id="M3"><mml:msubsup><mml:mrow><mml:mi>R</mml:mi></mml:mrow><mml:mrow><mml:mi>c</mml:mi></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:msubsup></mml:math></inline-formula> = 0.9845. Fan et al. (<xref ref-type="bibr" rid="B9">2014</xref>) adopted a combination of the standard normal variate, uninformative variable elimination, genetic algorithm, and successive projections algorithm to obtain 30 characteristic wavelengths selected from full-spectra achieving the optimal performance.</p>
<p>In the current study, the binary classification of Miyagawa Satsuma is focused on, which owns the regional characteristics of Chongming Island. The classification model for nondestructive determination of Miyagawa Satsuma SSC will judge the quality of citrus more scientifically, and overcome the shortcomings of subjective differences and low efficiency. Meanwhile, it can identify the growth cycle of citrus and estimate the maturity time more accurately, which is conducive to the management arrangement such as picking. It can also provide a theoretical basis for citrus grading, which is good for fruit farmers or manufacturers to sell graded citrus and improve profits (Kundu et al., <xref ref-type="bibr" rid="B16">2021</xref>).</p></sec>
<sec sec-type="materials and methods" id="s2">
<title>2. Materials and Methods</title>
<sec>
<title>2.1. Data Collection</title>
<sec>
<title>2.1.1. Near Infrared Spectra Acquisition</title>
<p>The equipment employed in our research is the Fourier transform near infrared spectrometer, an antaris II&#x02013;F-NIR analyzer made by Thermo Fisher. We set NIR acquisition mode as integrating sphere mode and the gain as &#x000D7;1. The NIR spectra are within the range from 1,000 nm to 2,500 nm.</p>
<p>All samples of Miyagawa Satsuma came from Shanghai Qianwei citrus Co., Ltd located on Chongming Island. All samplings, 11 times total, were carried out within 3 months. In each sampling, five trees with the most similar growth were selected, which were without films and not among the outermost three rows of trees. Then the five trees were divided into upper, middle, and lower parts, where one sample was picked, respectively, from four directions: south, east, north, and west. As a result, 12 samples were obtained per tree, and a total of 60 were taken for each sampling. Next, 12 samples were randomly chosen among a total of 60 fruits. For each sample in the 12 fruits, the NIR spectra were gained from 4 points at the cross symmetry of the equatorial plane of the fruit. Finally, the averaged NIR spectra, obtained by averaging NIR spectra of four points, were taken as the original NIR spectra, as shown in <xref ref-type="fig" rid="F1">Figure 1</xref>.</p>
<fig id="F1" position="float">
<label>Figure 1</label>
<caption><p>Near infrared (NIR) spectra of citrus in different picking times in chronological order. Different picking orders are represented by different colors. Each line is the average spectrum of 12 fruit samples at each picking time.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-13-841452-g0001.tif"/>
</fig></sec>
<sec>
<title>2.1.2. Soluble Solids Content Acquisition</title>
<p>After the Miyagawa Satsuma was squeezed and centrifuged, the SSC of the selected Miyagawa Satsuma samples was measured with a saccharometer (PR-101; Atago Co., Tokyo, Japan).</p></sec></sec>
<sec>
<title>2.2. Data Preprocessing</title>
<p>The samples with obviously incomplete or wrong data are eliminated, whether NIR spectra or SSC, thus obtaining a total of 122 samples. Then samples were divided into the training set and test set by the SPXY algorithm (Galvao et al., <xref ref-type="bibr" rid="B12">2005</xref>), with 70% of the samples as the training set and 30% as the test set. The principle of the Kennard stone algorithm (KS) algorithm is to calculate the Euclidean distance among all samples: select two samples with the maximum Euclidean distance into the training set, then carry out the iterative calculation, select the samples with the maximum and minimum Euclidean distance into the training set until the number of samples required by the training set is reached. SPXY algorithm is based on the KS algorithm, and it furthermore involves the chemical values and spectra among samples when calculating Euclidean distance, which makes the training set more representative, and makes the generalization ability of the established prediction model better.</p></sec>
<sec>
<title>2.3. Characteristic Wavelength Selection</title>
<p>The random frog (RF) algorithm (Li et al., <xref ref-type="bibr" rid="B17">2012</xref>) was used to obtain the corresponding number of characteristic wavelengths of NIR spectral data, which has the features of conceptually simplicity, and fewer parameters to be trained in algorithm implementation, strong global search and optimization ability, etc. The principles of the algorithm are as follows. Each sample in a population is regarded as a frog. Then the whole population is divided into <italic>m</italic> sub-groups with the scale of <italic>n</italic>. In each sub-group, the frogs with the best and worst fitness are used to produce a new child frog, which can be viewed as a jump of the best frog. If the fitness of the child frog is better than the parent with the worst fitness, replace the worst parent with the child, otherwise, randomly generate a new child, which can be viewed as the best frog&#x00027;s jumping again. If the fitness of the new child frog is still worse than the worst parent, then randomly generate another new child to replace the parent with the worst fitness. The evolutionary strategy of the random frog algorithm is like frogs jumping toward the optimal solution so that the algorithm gradually converges to the optimal solution.</p>
<p>The more specific steps of this algorithm are as follows: First, initialize parameters. Second, randomly generate an initial frog group and calculate the fitness of each frog. Third, arrange the frogs in descending order according to the value of fitness, and record the local optimal solution <italic>P</italic><sub><italic>x</italic></sub>. Then divide the <italic>F</italic> frogs from the initial group into sub-groups, namely, allocate <italic>F</italic> frogs into <italic>m</italic> sub-groups with the scale of <italic>n</italic>. Fourth, do a local search process, i.e., do the process described above in each sub-group. As a result, sub-groups do the fourth process, redivide the frog group, do the same operation as the first round, and record the global optimal solution <italic>P</italic><sub><italic>x</italic></sub>. Fifth, verify the calculation stop condition. If the convergence conditions of the algorithm are reached, the RF algorithm ends. If the global optimal solution has not been significantly improved, the execution of the algorithm should also be stopped.</p>
<p>To validate the performance of the RF algorithm in this task, the other common wavelength selection namely the competitive adaptive reweighted sampling algorithm (CARS) is used for comparison with the RF algorithm.</p></sec>
<sec>
<title>2.4. Binary Classification Model</title>
<p>The AdaBoost classifier (Freund et al., <xref ref-type="bibr" rid="B11">1999</xref>) is selected for modeling. Boosting is an important integrated learning technology, which can enhance weak classifiers with poor prediction performance into strong classifications with good prediction performance in a cascade way. The core of its adaptability is that the wrong samples of the previous basic classifier will be strengthened, and all the weighted samples will be used to train the next basic classifier again. At the same time, a new weak classifier is added in each round until a predetermined small enough error rate or a predetermined maximum number of iterations is reached.</p>
<p>Specifically, the entire AdaBoost iterative algorithm consists of three steps: First, initialize the weight distribution of training data. If there are <italic>n</italic> samples, each training sample is given the same weight of 1/<italic>n</italic> at the beginning. Second, train weak classifiers. In the specific training process, if a sample point has been accurately classified, its weight will be reduced in the construction of the next training set; On the contrary, if a sample point is not accurately classified, its weight will be improved. Then, the weight updated sample set is used to train the next classifier, and the whole training process goes on in this way, iteratively. Third, combine the trained weak classifiers into strong classifiers. After the training process of each weak classifier, increase the weight of the weak classifier with a small classification error rate to make it play a greater decisive role in the final classification function while doing the opposite operation for the weak classifier with a large classification error rate.</p>
<p>To compare the performance of various classifiers, we choose AdaBoost, k-Nearest Neighbor (KNN), Bayes classifier, and LS-SVM to explore the best-performing classification model. In the current study, we use Matlab and Weka to establish the models.</p></sec>
<sec>
<title>2.5. Model Evaluation</title>
<p>To verify the efficiency of the classification system, evaluation indicators viz. confusion matrix, accuracy, precision, recall, <italic>F</italic><sub>1</sub>, micro-measures, and macro-measures are considered.</p>
<p>1) Confusion matrix: Assume that &#x0201C;Positive&#x0201D; means the positive samples and that &#x0201C;Negative&#x0201D; means the negative samples. Meanwhile, &#x0201C;True&#x0201D; represents that the prediction is right while &#x0201C;False&#x0201D; represents that the prediction is wrong. As a result, &#x0201C;TP&#x0201D; and &#x0201C;TN&#x0201D; mean that the positive sample is classified as &#x0201C;Positive&#x0201D; and that the negative sample is labeled as &#x0201C;Negative&#x0201D;, respectively. &#x0201C;FP&#x0201D; and &#x0201C;FN&#x0201D; represent that the negative sample is labeled as &#x0201C;Positive&#x0201D; and that the positive sample is classified as &#x0201C;False.&#x0201D; The four indicators make up the confusion matrix.</p>
<p>2) Accuracy: It is a ratio that is used to estimate the classification ability of a model within the range from 0 to 1. Generally speaking, the larger accuracy is, the better the classification is. It can be calculated by the following equation:</p>
<disp-formula id="E1"><label>(1)</label><mml:math id="M4"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>A</mml:mi><mml:mi>c</mml:mi><mml:mi>c</mml:mi><mml:mi>u</mml:mi><mml:mi>r</mml:mi><mml:mi>a</mml:mi><mml:mi>c</mml:mi><mml:mi>y</mml:mi><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>N</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>T</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>N</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>N</mml:mi></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>3) Precision: Precision is only used to evaluate the classification ability of the positive samples within the range from 0 to 1. It is obvious that the larger precision is, the more effective the system is. It is computed by:</p>
<disp-formula id="E2"><label>(2)</label><mml:math id="M5"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>P</mml:mi><mml:mi>r</mml:mi><mml:mi>e</mml:mi><mml:mi>c</mml:mi><mml:mi>i</mml:mi><mml:mi>s</mml:mi><mml:mi>i</mml:mi><mml:mi>o</mml:mi><mml:mi>n</mml:mi><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>P</mml:mi></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>4) Recall: It is a ratio from 0 to 1. Obviously, the more it is close to 1, the better the system is. The calculation equation is:</p>
<disp-formula id="E3"><label>(3)</label><mml:math id="M6"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>R</mml:mi><mml:mi>e</mml:mi><mml:mi>c</mml:mi><mml:mi>a</mml:mi><mml:mi>l</mml:mi><mml:mi>l</mml:mi><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>N</mml:mi></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>5) <italic>F</italic><sub>1</sub>: It is a harmonic mean of recall and precision. In this study, we consider the weight of recall and precision the same, which means attaching the weight of 0.5 to either of them. It is calculated by:</p>
<disp-formula id="E4"><label>(4)</label><mml:math id="M7"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:msub><mml:mrow><mml:mi>F</mml:mi></mml:mrow><mml:mrow><mml:mn>1</mml:mn></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mn>2</mml:mn><mml:mo>&#x0002A;</mml:mo><mml:mi>P</mml:mi><mml:mi>r</mml:mi><mml:mi>e</mml:mi><mml:mi>c</mml:mi><mml:mi>i</mml:mi><mml:mi>s</mml:mi><mml:mi>i</mml:mi><mml:mi>o</mml:mi><mml:mi>n</mml:mi><mml:mo>&#x0002A;</mml:mo><mml:mi>R</mml:mi><mml:mi>e</mml:mi><mml:mi>c</mml:mi><mml:mi>a</mml:mi><mml:mi>l</mml:mi><mml:mi>l</mml:mi></mml:mrow><mml:mrow><mml:mi>P</mml:mi><mml:mi>r</mml:mi><mml:mi>e</mml:mi><mml:mi>c</mml:mi><mml:mi>i</mml:mi><mml:mi>s</mml:mi><mml:mi>i</mml:mi><mml:mi>o</mml:mi><mml:mi>n</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>R</mml:mi><mml:mi>e</mml:mi><mml:mi>c</mml:mi><mml:mi>a</mml:mi><mml:mi>l</mml:mi><mml:mi>l</mml:mi></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>6) Receiver Operating Characteristic (ROC) Curve and Area Under Curve (AUC): The abscissa of the ROC curve is the false positive rate (FPR) while the ordinate is the true positive rate (TPR), where <inline-formula><mml:math id="M8"><mml:mtext>FPR</mml:mtext><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>F</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>N</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>P</mml:mi></mml:mrow></mml:mfrac></mml:math></inline-formula> and <inline-formula><mml:math id="M9"><mml:mtext>TPR</mml:mtext><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>Y</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>N</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>N</mml:mi></mml:mrow></mml:mfrac></mml:math></inline-formula>. Generally speaking, the closer the ROC curve is to the upper left corner of the image, the better the performance of the binary classifier. AUC is the area under the ROC curve and it is generally within the range of (0.5, 1). When the closer the ROC curve is to the upper left corner, the greater the value of AUC.</p></sec></sec>
<sec id="s3">
<title>3. Experimental Results and Analysis</title>
<sec>
<title>3.1. NIR Spectral Characteristics of Miyagawa Satsuma in Different Picking Time</title>
<p>The NIR spectra in different picking times are shown in <xref ref-type="fig" rid="F1">Figure 1</xref>. The trend of the citrus spectra collected each time is similar. There is an obvious absorption trough near 1,080, 1,300, 1,700, and 2,200 nm, respectively. According to the principles of NIR spectroscopy, due to the fact that the sample will selectively absorb NIR waves with different frequencies, the NIR wave which passes through the sample will become weaker in some wavelength ranges, and the transmitted NIR wave will carry the information of organic component and structure. Therefore, it can be inferred that these absorption troughs can probably be the most informative areas, which can be reflected in the characteristic wavelength selection.</p>
<p>Our reason for using fruits with different picking periods for modeling is to increase the coverage of the SSC, allowing a larger range of variation in the spectral data and ultimately increasing the model robustness. We performed the statistical tests on the obtained spectral data and SSC and found significant differences between spectral data and SSC for non-adjacent picking periods (<italic>p</italic> &#x0003C; 0.05) and no significant differences for adjacent picking periods (<italic>p</italic> &#x0003E; 0.05). This is in accordance with expectations. Because, as the fruit ripens, the SSC will certainly increase and the spectral differences will increase.</p></sec>
<sec>
<title>3.2. Performance of RF</title>
<p>As mentioned above, the characteristic wavelength selection can accelerate the computation speed and reduce computation cost to a degree. RF algorithm is chosen to generate characteristic wavelengths with the numbers 10, 50, 100, 200, 250, 275, and 300, respectively, which is displayed in <xref ref-type="fig" rid="F2">Figure 2</xref>. It is easy to find that the larger the number of characteristic wavelengths is, the smaller the cutoff probability is. The cutoff probability indicates the threshold value for screening the required number of informative wavelengths. The wavelength numbers, 10, 50, 100, 200, 250, 275, and 300, respectively correspond to the cutoff probabilities, 0.0471, 0.0260, 0.0225, 0.0222, 0.0091, 0.0060, and 0.0047. The cutoff probability generally decreases as the number of informative wavelengths increases (<xref ref-type="fig" rid="F2">Figure 2</xref>). Meanwhile, it is true with what has been inferred in the above section that the absorption troughs can be the most informative, most of the retaining wavelengths gather in the areas inferred before viz. 1,080, 1,300, 1,700, and 2,200 nm. This probably has a relationship with the functional groups viz. &#x02014;OH, &#x02014;CH, &#x02014;NH.</p>
<fig id="F2" position="float">
<label>Figure 2</label>
<caption><p>Cutoff probability of different characteristic wavelengths. For the subsequent modeling, 10 <bold>(A)</bold>, 50 <bold>(B)</bold>, 100 <bold>(C)</bold>, 200 <bold>(D)</bold>, 250 <bold>(E)</bold>, 275 <bold>(F)</bold>, and 300 <bold>(G)</bold> characteristic wavelengths are chosen, respectively. The numbers on the right corner of the figure are the cutoff probability of corresponding characteristic wavelengths number. The cutoff probability generally decreases as the number of informative wavelengths increases.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-13-841452-g0002.tif"/>
</fig>
<p>In addition, the classification models based on CARS selected wavelengths are established, and their performance is not as good as the RF-based models. For example, when ten characteristic wavelengths are selected, the RF-based model gives a better performance than the model based on CARS, with the accuracy of 60.9, 69.6, 65.2, and 62.2% vs. 52.78, 52.78, 58.33, and 47.22% for AdaBoost, KNN, Bayes, and LS-SVM modeling methods, respectively. Overall, CARS does not perform as well as RF for the informative wavelength selection.</p></sec>
<sec>
<title>3.3. Model Analysis Using Plant Physiology Phenomenon</title>
<p>The spectral properties of plants are mainly determined by their internal structure. For the current study, the obtained spectra are the result of the interaction of the incident light with the chemical composition and physical structure of Citrus. For Miyagawa Satsuma, its structure can be divided into exocarp (oil cell layer), mesocarp (white cortex), endocarp, fruit, and fruit stem from outside to inside. Among them, the surface of the soluble dietary fiber of mandarin pulp is not smooth, the strips and gaps are intertwined, and there are raised particles; the surface of the soluble dietary fiber of mandarin peel is larger, but the surface depressions are mixed with a few spherical particles. There is a strong interaction between the two molecules.</p>
<p>As an important indicator for evaluating fruit sweetness, SSC is mainly composed of soluble sugars (including sucrose, fructose, and glucose). In the NIR region, the stretching and deformation vibration absorption peaks of <italic>O</italic>&#x02212;<italic>H</italic> bonds in soluble sugars are located around 1,440 and 2,080 nm, and there are three absorption peaks of soluble solids at 980, 1,169, and 1,485 nm (Musingarabwi et al., <xref ref-type="bibr" rid="B26">2016</xref>). The water content has a great influence on the absorption of the plant spectrum. Under the condition of multi-layer leaves, the water absorption bands at 1,100 and 960 nm have a great influence on spectral reflectance. Absorption leads to a decrease in reflectance and an increase in absorbance, and peaks of reflectance (i.e., peaks and valleys of absorbance) appear at 1,600 and 2,200 nm (Ma et al., <xref ref-type="bibr" rid="B23">2017</xref>).</p>
<p>The wavelengths selected by RF include three characteristic wavelengths near the absorption peaks of soluble solids at 980 nm, 1,169 nm, and 1,485 nm, and two characteristic wavelengths near the absorption peaks of stretching and deformation vibrations of <italic>O</italic>&#x02212;<italic>H</italic> bonds in soluble sugars at 1,440 nm and 2,080 nm, and a characteristic wavelength near the strong absorption peak of water at 1,400 nm. This analysis explains why the model based on the RF selected wavelengths performs better.</p></sec>
<sec>
<title>3.4. Soluble Solids Content Division</title>
<p>The research holds the opinion that consumers are more concerned about whether the Miyagawa Satsuma is sweet or not, but not the concrete value of sweetness. Referring to <xref ref-type="fig" rid="F3">Figure 3</xref>, the dichotomous map or histogram of 122 Miyagawa Satsuma citruses&#x00027; SSC, the distribution of this figure is roughly similar to the normal distribution, whose mean of all citruses&#x00027; SSC is 9.06 Brix and the SD is 0.93 Brix. To carry out our belief, 9 Brix was taken as the boundary after asking an expert in agriculture for advice. As a consequence, the citruses with SSC more than or equal to 9 Brix are considered to be sweet and the others are not sweet for the following classification modeling.</p>
<fig id="F3" position="float">
<label>Figure 3</label>
<caption><p>Histogram of citrus soluble solid content (SSC) of the total 122 fruit samples collected in this study. The mean of all citruses&#x00027; SSC is 9.06 Brix while the SD is 0.93 Brix. The citruses corresponding to the blue areas are regarded as sweet while the citruses corresponding to the red ones are regarded as unsweet.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-13-841452-g0003.tif"/>
</fig>
<p><xref ref-type="fig" rid="F4">Figure 4</xref> shows the comparison of NIR spectra of sweet and unsweet fruit samples. As shown in <xref ref-type="fig" rid="F4">Figure 4</xref>, the large overlap between the spectral curves of the sweet and unsweet samples indicates that the model will not perform as expected if the model is constructed based on original spectra. Therefore, we need to select the informative wavelengths specific to SSC classification, and then combine them with pattern recognition methods for modeling.</p>
<fig id="F4" position="float">
<label>Figure 4</label>
<caption><p>Comparison of NIR spectra of sweet (SSC beyond 9 Brix) and unsweet (SSC below 9 Brix) fruit samples. The shaded areas represent the confident intervals to each line.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-13-841452-g0004.tif"/>
</fig></sec>
<sec>
<title>3.5. Performance Analysis of Different Binary Classification Models</title>
<p>As mentioned before, AdaBoost, KNN, Bayes, and LS-SVM are adopted to establish classification models. The performance comparison of different classifiers under different characteristic wavelengths is shown in <xref ref-type="table" rid="T1">Table 1</xref>. The conclusion can be drawn that when the number of characteristic wavelengths is 275, the classification model established by the AdaBoost classifier performs best (bold in the table), with accuracy, precision, recall, and <italic>F</italic><sub>1</sub>-score 78.3%, 80.5%, 78.3%, 0.780, respectively.</p>
<table-wrap position="float" id="T1">
<label>Table 1</label>
<caption><p>Modeling results of sweet (SSC beyond 9 Brix) and unsweet (SSC below 9 Brix) classification of Miyagawa Satsuma in Chongming Island under different classifiers with different characteristic wavelengths.</p></caption>
<table frame="hsides" rules="groups">
<thead><tr>
<th valign="top" align="left"><bold>Characteristic wavelengths</bold></th>
<th valign="top" align="left"><bold>Models</bold></th>
<th/>
<th valign="top" align="center"><bold>Metrics</bold></th>
<th/>
<th/>
</tr>
<tr>
<th/>
<th/>
<th valign="top" align="center"><bold>Accuracy</bold></th>
<th valign="top" align="center"><bold>Precision</bold></th>
<th valign="top" align="center"><bold>Recall</bold></th>
<th valign="top" align="center"><bold>F1-score</bold></th>
</tr>
<tr>
<th/>
<th/>
<th valign="top" align="center"><bold>(%)</bold></th>
<th valign="top" align="center"><bold>(%)</bold></th>
<th valign="top" align="center"><bold>(%)</bold></th>
<th/>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">10</td>
<td valign="top" align="left">AdaBoost</td>
<td valign="top" align="center">60.9</td>
<td valign="top" align="center">62.1</td>
<td valign="top" align="center">60.9</td>
<td valign="top" align="center">0.604</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">KNN</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.8</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">Bayes</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">65.4</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">0.648</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">LS-SVM</td>
<td valign="top" align="center">62.2</td>
<td valign="top" align="center">53.3</td>
<td valign="top" align="center">100.0</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td valign="top" align="left">50</td>
<td valign="top" align="left">AdaBoost</td>
<td valign="top" align="center">60.9</td>
<td valign="top" align="center">62.1</td>
<td valign="top" align="center">60.9</td>
<td valign="top" align="center">0.604</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">KNN</td>
<td valign="top" align="center">52.2</td>
<td valign="top" align="center">53.7</td>
<td valign="top" align="center">52.2</td>
<td valign="top" align="center">0.503</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">Bayes</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">65.4</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">0.648</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">LS-SVM</td>
<td valign="top" align="center">67.6</td>
<td valign="top" align="center">57.1</td>
<td valign="top" align="center">100.0</td>
<td valign="top" align="center">0.727</td>
</tr>
<tr>
<td valign="top" align="left">100</td>
<td valign="top" align="left">AdaBoost</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.8</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">KNN</td>
<td valign="top" align="center">52.2</td>
<td valign="top" align="center">52.9</td>
<td valign="top" align="center">52.2</td>
<td valign="top" align="center">0.516</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">Bayes</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">0.694</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">LS-SVM</td>
<td valign="top" align="center">59.5</td>
<td valign="top" align="center">51.6</td>
<td valign="top" align="center">100.0</td>
<td valign="top" align="center">0.681</td>
</tr>
<tr>
<td valign="top" align="left">200</td>
<td valign="top" align="left">AdaBoost</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">71.6</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">0.631</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">KNN</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">67.8</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">0.644</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">Bayes</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">0.694</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">LS-SVM</td>
<td valign="top" align="center">62.2</td>
<td valign="top" align="center">53.3</td>
<td valign="top" align="center">100.0</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td valign="top" align="left">250</td>
<td valign="top" align="left">AdaBoost</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">71.3</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">0.692</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">KNN</td>
<td valign="top" align="center">56.5</td>
<td valign="top" align="center">58.1</td>
<td valign="top" align="center">56.5</td>
<td valign="top" align="center">0.555</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">Bayes</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.8</td>
<td valign="top" align="center">0.7</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">LS-SVM</td>
<td valign="top" align="center">62.2</td>
<td valign="top" align="center">53.3</td>
<td valign="top" align="center">100.0</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td valign="top" align="left">275</td>
<td valign="top" align="left"><bold>AdaBoost</bold></td>
<td valign="top" align="center"><bold>78.3</bold></td>
<td valign="top" align="center"><bold>80.5</bold></td>
<td valign="top" align="center"><bold>78.3</bold></td>
<td valign="top" align="center"><bold>0.780</bold></td>
</tr>
<tr>
<td/>
<td valign="top" align="left">KNN</td>
<td valign="top" align="center">60.9</td>
<td valign="top" align="center">62.1</td>
<td valign="top" align="center">60.9</td>
<td valign="top" align="center">0.604</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">Bayes</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.8</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">LS-SVM</td>
<td valign="top" align="center">62.2</td>
<td valign="top" align="center">53.3</td>
<td valign="top" align="center">100.0</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td valign="top" align="left">300</td>
<td valign="top" align="left">AdaBoost</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">71.3</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">0.692</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">KNN</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">67.8</td>
<td valign="top" align="center">65.2</td>
<td valign="top" align="center">0.644</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">Bayes</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">69.8</td>
<td valign="top" align="center">69.6</td>
<td valign="top" align="center">0.696</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">LS-SVM</td>
<td valign="top" align="center">62.2</td>
<td valign="top" align="center">54.2</td>
<td valign="top" align="center">81.3</td>
<td valign="top" align="center">0.65</td>
</tr>
<tr>
<td valign="top" align="left">1556</td>
<td valign="top" align="left">AdaBoost</td>
<td valign="top" align="center">75.0</td>
<td valign="top" align="center">68.8</td>
<td valign="top" align="center">91.7</td>
<td valign="top" align="center">0.786</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">KNN</td>
<td valign="top" align="center">60.9</td>
<td valign="top" align="center">56.3</td>
<td valign="top" align="center">81.8</td>
<td valign="top" align="center">0.667</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">Bayes</td>
<td valign="top" align="center">73.9</td>
<td valign="top" align="center">72.7</td>
<td valign="top" align="center">72.7</td>
<td valign="top" align="center">0.727</td>
</tr>
<tr>
<td/>
<td valign="top" align="left">LS-SVM</td>
<td valign="top" align="center">56.8</td>
<td valign="top" align="center">50.0</td>
<td valign="top" align="center">93.8</td>
<td valign="top" align="center">0.652</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p><italic>Bold font represents the best model</italic>.</p>
</table-wrap-foot>
</table-wrap>
<p>From the perspective of the number of characteristic wavelengths, when the number is 10, the best performer is the KNN classifier, with accuracy, precision, recall, and <italic>F</italic><sub>1</sub>-score 69.6%, 69.8%, 69.6%, and 0.696. When the number is 50, LS-SVM performs best according to the accuracy of 67.6%, precision of 57.1%, recall of 100%, and <italic>F</italic><sub>1</sub>-score 0.727. When the number is 100, Adaboost performs best while the best performer belongs to Bayes when the number is 200. As for the number 250, the results of AdaBoost are as good as Bayes. Finally, AdaBoost still stands out among four classifiers when under the condition of 300 characteristic wavelengths. From <xref ref-type="table" rid="T1">Table 1</xref>, it can be found that when the number of characteristic wavelengths is either too small or too large, the performance of every different classifier is not as good as the situation when the number is proper, from the perspective of four classifiers.</p>
<p>Compared to the results of original wavelengths number 1,556 without any procession, the best results of AdaBoost, KNN, and LS-SVM happen when they are through characteristic wavelengths selection, however, except Bayes. But after weighing the wavelength reduction and performance, it is reasonable to think that characteristic wavelength selection also works for Bayes.</p>
<p>The ROC of positive and negative samples of the test set is shown in <xref ref-type="fig" rid="F5">Figure 5</xref>. It can be seen that the ROC curves of positive and negative samples are all close to the upper left corner, and the total AUC is 0.841, indicating that the model has good robustness and can adapt flexibly to the uneven distribution of positive and negative samples in actual situations.</p>
<fig id="F5" position="float">
<label>Figure 5</label>
<caption><p>Receiver Operating Characteristic (ROC) curves. <bold>(A)</bold> is ROC curve of positive samples (SSC beyond 9 Brix) while <bold>(B)</bold> is ROC curve of negative samples (SSC below 9 Brix).</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-13-841452-g0005.tif"/>
</fig>
<p>Too many spectral features bring information redundancy, and too few spectral features bring information loss. Based on the experimental results, for this classification task, the optimal number of spectral features is 275. Compared to the other modeling methods, the AdaBoost method achieves the best performance at 275 wavelength numbers. This is because AdaBoost combines multiple weak classifiers in a reasonable way to make one strong classifier. The other three methods used in this paper just give one separate model.</p></sec></sec>
<sec id="s4">
<title>4. Conclusion and Reflection</title>
<p>Based on NIR spectroscopy, the random frog algorithm, and AdaBoost algorithm, and taking citrus in Shanghai Chongming Island as the research object, this study focuses on the problems of binary classification between NIR spectra and Miyagawa Satsuma SSC. Nine Brix is selected as the threshold of being sweet or not and the samples are divided into the training set and test set. After selecting characteristic wavelengths through the RF algorithm, they are used to establish binary classification models by AdaBoost, LS-SVM, and other classifiers. According to their performance, the AdaBoost classifier is the optimum model, with accuracy, precision, recall, and <italic>F</italic><sub>1</sub>-score 78.3%, 80.5%, 78.3%, and 0.780, respectively.</p>
<p>Analyzing the model performance, we find that the constructed model does not have a very high performance. Combined with the sampling process and the test results, two reasons may be summarized (1) due to the limited penetration depth of NIR and the thick skin of the fruit, most of the NIR light does not penetrate the skin to reach the fruit part; and (2) there are environment disturbances during sampling and instrument errors in the process of collecting spectra.</p>
<p>The constructed model has the potential to be embedded in portable NIR acquisition devices in the future, which can facilitate fruit farmers to judge the quality of the citrus and be conducive to improving the sale pricing system of citrus in Chongming Island, so as to maximize the sale profit of fruit-sellers.</p></sec>
<sec sec-type="data-availability" id="s5">
<title>Data Availability Statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p></sec>
<sec id="s6">
<title>Author Contributions</title>
<p>CZ and MH: funding acquisition. YC, MH, WS, SJ, LW, BD, ZC, and CZ: methodology and validation. YC, WS, and MH: writing&#x02014;original draft. YC, SJ, LW, MH, and CZ: writing&#x02014;review and editing. FJ: providing citrus materials. All authors contributed to the article and approved the submitted version.</p></sec>
<sec sec-type="funding-information" id="s7">
<title>Funding</title>
<p>This study is sponsored by Shanghai Agriculture Applied Technology Development Program (Grant No. X20200102), Shanghai Agricultural System Standard Development Program (2018-013), Agriculture Research System of Shanghai (Grant No. 201407), Shanghai Agriculture Applied Technology Development Program (Grant No. T20220103), the Science and Technology Commission of Shanghai Municipality (No. 19511120100), the GHfund B (No. 20210702), and the Fundamental Research Funds for the Central Universities.</p></sec>
<sec sec-type="COI-statement" id="conf1">
<title>Conflict of Interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p></sec>
<sec sec-type="disclaimer" id="s8">
<title>Publisher&#x00027;s Note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p></sec>
</body>
<back>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Anticona</surname> <given-names>M.</given-names></name> <name><surname>Blesa</surname> <given-names>J.</given-names></name> <name><surname>Frigola</surname> <given-names>A.</given-names></name> <name><surname>Esteve</surname> <given-names>M. J.</given-names></name></person-group> (<year>2020</year>). <article-title>High biological value compounds extraction from citrus waste with non-conventional methods</article-title>. <source>Foods</source> <volume>9</volume>, <fpage>811</fpage>. <pub-id pub-id-type="doi">10.3390/foods9060811</pub-id><pub-id pub-id-type="pmid">32575685</pub-id></citation></ref>
<ref id="B2">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Arefi</surname> <given-names>A.</given-names></name> <name><surname>Sturm</surname> <given-names>B.</given-names></name> <name><surname>von Gersdorff</surname> <given-names>G.</given-names></name> <name><surname>Nasirahmadi</surname> <given-names>A.</given-names></name> <name><surname>Hensel</surname> <given-names>O.</given-names></name></person-group> (<year>2021</year>). <article-title>Vis-nir hyperspectral imaging along with gaussian process regression to monitor quality attributes of apple slices during drying</article-title>. <source>LWT</source> <volume>152</volume>, <fpage>112297</fpage>. <pub-id pub-id-type="doi">10.1016/j.lwt.2021.112297</pub-id></citation>
</ref>
<ref id="B3">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Arendse</surname> <given-names>E.</given-names></name> <name><surname>Fawole</surname> <given-names>O. A.</given-names></name> <name><surname>Magwaza</surname> <given-names>L. S.</given-names></name> <name><surname>Opara</surname> <given-names>U. L.</given-names></name></person-group> (<year>2018</year>). <article-title>Non-destructive prediction of internal and external quality attributes of fruit with thick rind: a review</article-title>. <source>J. Food Eng</source>. <volume>217</volume>, <fpage>11</fpage>&#x02013;<lpage>23</lpage>. <pub-id pub-id-type="doi">10.1016/j.jfoodeng.2017.08.009</pub-id></citation>
</ref>
<ref id="B4">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Beghi</surname> <given-names>R.</given-names></name> <name><surname>Buratti</surname> <given-names>S.</given-names></name> <name><surname>Giovenzana</surname> <given-names>V.</given-names></name> <name><surname>Benedetti</surname> <given-names>S.</given-names></name> <name><surname>Guidetti</surname> <given-names>R.</given-names></name></person-group> (<year>2017</year>). <article-title>Electronic nose and visible-near infrared spectroscopy in fruit and vegetable monitoring</article-title>. <source>Rev. Anal. Chem</source>. <volume>36</volume>, <fpage>1</fpage>&#x02013;<lpage>24</lpage>. <pub-id pub-id-type="doi">10.1515/revac-2016-0016</pub-id></citation>
</ref>
<ref id="B5">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bhargava</surname> <given-names>A.</given-names></name> <name><surname>Bansal</surname> <given-names>A.</given-names></name></person-group> (<year>2021</year>). <article-title>Fruits and vegetables quality evaluation using computer vision: a review</article-title>. <source>J. King Saud Univer. Comput. Inf. Sci</source>. <volume>33</volume>, <fpage>243</fpage>&#x02013;<lpage>257</lpage>. <pub-id pub-id-type="doi">10.1016/j.jksuci.2018.06.002</pub-id><pub-id pub-id-type="pmid">24915395</pub-id></citation></ref>
<ref id="B6">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chang</surname> <given-names>H.</given-names></name> <name><surname>Zhu</surname> <given-names>L.</given-names></name> <name><surname>Lou</surname> <given-names>X.</given-names></name> <name><surname>Meng</surname> <given-names>X.</given-names></name> <name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Wang</surname> <given-names>Z.</given-names></name></person-group> (<year>2016</year>). <article-title>Local strategy combined with a wavelength selection method for multivariate calibration</article-title>. <source>Sensors</source> <volume>16</volume>, <fpage>827</fpage>. <pub-id pub-id-type="doi">10.3390/s16060827</pub-id><pub-id pub-id-type="pmid">27271636</pub-id></citation></ref>
<ref id="B7">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Cheng</surname> <given-names>C.-,x.</given-names></name> <name><surname>Jia</surname> <given-names>M.</given-names></name> <name><surname>Gui</surname> <given-names>Y.</given-names></name> <name><surname>Ma</surname> <given-names>Y.</given-names></name></person-group> (<year>2020</year>). <article-title>Comparison of the effects of novel processing technologies and conventional thermal pasteurisation on the nutritional quality and aroma of mandarin (citrus unshiu) juice</article-title>. <source>Innovat. Food Sci. Emerg. Technol</source>. 64, 102425. <pub-id pub-id-type="doi">10.1016/j.ifset.2020.102425</pub-id></citation>
</ref>
<ref id="B8">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Cruz</surname> <given-names>S.</given-names></name> <name><surname>Guerra</surname> <given-names>R.</given-names></name> <name><surname>Brazio</surname> <given-names>A.</given-names></name> <name><surname>Cavaco</surname> <given-names>A. M.</given-names></name> <name><surname>Antunes</surname> <given-names>D.</given-names></name> <name><surname>Passos</surname> <given-names>D.</given-names></name></person-group> (<year>2021</year>). <article-title>Nondestructive simultaneous prediction of internal browning disorder and quality attributes in &#x02018;rocha&#x00027;pear (<italic>Pyrus communis</italic> L.) using vis-nir spectroscopy</article-title>. <source>Postharvest Biol. Technol</source>. 179, 111562. <pub-id pub-id-type="doi">10.1016/j.postharvbio.2021.111562</pub-id></citation>
</ref>
<ref id="B9">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fan</surname> <given-names>S.</given-names></name> <name><surname>Huang</surname> <given-names>W.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Zhao</surname> <given-names>C.</given-names></name> <name><surname>Zhang</surname> <given-names>B.</given-names></name></person-group> (<year>2014</year>). <article-title>Characteristic wavelengths selection of soluble solids content of pear based on nir spectral and ls-svm</article-title>. <source>Guang pu xue yu guang pu fen xi= Guang pu</source> <volume>34</volume>, <fpage>2089</fpage>&#x02013;<lpage>2093</lpage>. <pub-id pub-id-type="doi">10.3964/j.issn.1000-0593(2014)08-2089-05</pub-id><pub-id pub-id-type="pmid">25474940</pub-id></citation></ref>
<ref id="B10">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fan</surname> <given-names>S.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Xia</surname> <given-names>Y.</given-names></name> <name><surname>Tian</surname> <given-names>X.</given-names></name> <name><surname>Guo</surname> <given-names>Z.</given-names></name> <name><surname>Huang</surname> <given-names>W.</given-names></name></person-group> (<year>2019</year>). <article-title>Long-term evaluation of soluble solids content of apples with biological variability by using near-infrared spectroscopy and calibration transfer method</article-title>. <source>Postharvest Biol. Technol</source>. <volume>151</volume>, <fpage>79</fpage>&#x02013;<lpage>87</lpage>. <pub-id pub-id-type="doi">10.1016/j.postharvbio.2019.02.001</pub-id></citation>
</ref>
<ref id="B11">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Freund</surname> <given-names>Y.</given-names></name> <name><surname>Schapire</surname> <given-names>R.</given-names></name> <name><surname>Abe</surname> <given-names>N.</given-names></name></person-group> (<year>1999</year>). <article-title>A short introduction to boosting</article-title>. <source>J. Jpn. Soc. Artif. Intell</source>. 14, 1612.</citation>
</ref>
<ref id="B12">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Galvao</surname> <given-names>R. K. H.</given-names></name> <name><surname>Araujo</surname> <given-names>M. C. U.</given-names></name> <name><surname>Jos&#x000E9;</surname> <given-names>G. E.</given-names></name> <name><surname>Pontes</surname> <given-names>M. J. C.</given-names></name> <name><surname>Silva</surname> <given-names>E. C.</given-names></name> <name><surname>Saldanha</surname> <given-names>T. C. B.</given-names></name></person-group> (<year>2005</year>). <article-title>A method for calibration and validation subset partitioning</article-title>. <source>Talanta</source> <volume>67</volume>, <fpage>736</fpage>&#x02013;<lpage>740</lpage>. <pub-id pub-id-type="doi">10.1016/j.talanta.2005.03.025</pub-id><pub-id pub-id-type="pmid">18970233</pub-id></citation></ref>
<ref id="B13">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>Z.</given-names></name> <name><surname>Wang</surname> <given-names>M.</given-names></name> <name><surname>Agyekum</surname> <given-names>A. A.</given-names></name> <name><surname>Wu</surname> <given-names>J.</given-names></name> <name><surname>Chen</surname> <given-names>Q.</given-names></name> <name><surname>Zuo</surname> <given-names>M.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>Quantitative detection of apple watercore and soluble solids content by near infrared transmittance spectroscopy</article-title>. <source>J. Food Eng</source>. 279, 109955. <pub-id pub-id-type="doi">10.1016/j.jfoodeng.2020.109955</pub-id></citation>
</ref>
<ref id="B14">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Huang</surname> <given-names>Y.</given-names></name> <name><surname>Dong</surname> <given-names>W.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Wang</surname> <given-names>X.</given-names></name> <name><surname>Luo</surname> <given-names>W.</given-names></name> <name><surname>Zhan</surname> <given-names>B.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Online detection of soluble solids content and maturity of tomatoes using vis/nir full transmittance spectra</article-title>. <source>Chemometr. Intell. Lab. Syst</source>. 210, 104243. <pub-id pub-id-type="doi">10.1016/j.chemolab.2021.104243</pub-id></citation>
</ref>
<ref id="B15">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Jun</surname> <given-names>S.</given-names></name> <name><surname>Yating</surname> <given-names>L.</given-names></name> <name><surname>Xiaohong</surname> <given-names>W.</given-names></name> <name><surname>Chunxia</surname> <given-names>D.</given-names></name> <name><surname>Yong</surname> <given-names>C.</given-names></name></person-group> (<year>2018</year>). <article-title>Ssc prediction of cherry tomatoes based on iriv-cs-svr model and near infrared reflectance spectroscopy</article-title>. <source>J. Food Process. Eng</source>. 41, e12884. <pub-id pub-id-type="doi">10.1111/jfpe.12884</pub-id></citation>
</ref>
<ref id="B16">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kundu</surname> <given-names>N.</given-names></name> <name><surname>Rani</surname> <given-names>G.</given-names></name> <name><surname>Dhaka</surname> <given-names>V. S.</given-names></name> <name><surname>Gupta</surname> <given-names>K.</given-names></name> <name><surname>Nayak</surname> <given-names>S. C.</given-names></name> <name><surname>Verma</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Iot and interpretable machine learning based framework for disease prediction in pearl millet</article-title>. <source>Sensors</source> <volume>21</volume>, <fpage>5386</fpage>. <pub-id pub-id-type="doi">10.3390/s21165386</pub-id><pub-id pub-id-type="pmid">34450827</pub-id></citation></ref>
<ref id="B17">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>H.-D.</given-names></name> <name><surname>Xu</surname> <given-names>Q.-S.</given-names></name> <name><surname>Liang</surname> <given-names>Y.-Z.</given-names></name></person-group> (<year>2012</year>). <article-title>Random frog: an efficient reversible jump markov chain monte carlo-like approach for variable selection with applications to gene selection and disease classification</article-title>. <source>Anal. Chim. Acta</source> <volume>740</volume>, <fpage>20</fpage>&#x02013;<lpage>26</lpage>. <pub-id pub-id-type="doi">10.1016/j.aca.2012.06.031</pub-id><pub-id pub-id-type="pmid">22840646</pub-id></citation></ref>
<ref id="B18">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Tian</surname> <given-names>X.</given-names></name> <name><surname>Huang</surname> <given-names>W.</given-names></name> <name><surname>Zhang</surname> <given-names>B.</given-names></name> <name><surname>Fan</surname> <given-names>S.</given-names></name></person-group> (<year>2016a</year>). <article-title>Application of long-wave near infrared hyperspectral imaging for measurement of soluble solid content (ssc) in pear</article-title>. <source>Food Anal. Methods</source> <volume>9</volume>, <fpage>3087</fpage>&#x02013;<lpage>3098</lpage>. <pub-id pub-id-type="doi">10.1007/s12161-016-0498-2</pub-id></citation>
</ref>
<ref id="B19">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Wang</surname> <given-names>Q.</given-names></name> <name><surname>Xu</surname> <given-names>L.</given-names></name> <name><surname>Tian</surname> <given-names>X.</given-names></name> <name><surname>Xia</surname> <given-names>Y.</given-names></name> <name><surname>Fan</surname> <given-names>S.</given-names></name></person-group> (<year>2019</year>). <article-title>Comparison and optimization of models for determination of sugar content in pear by portable vis-nir spectroscopy coupled with wavelength selection algorithm</article-title>. <source>Food Anal. Methods</source> <volume>12</volume>, <fpage>12</fpage>&#x02013;<lpage>22</lpage>. <pub-id pub-id-type="doi">10.1007/s12161-018-1326-7</pub-id></citation>
</ref>
<ref id="B20">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>J.-L.</given-names></name> <name><surname>Sun</surname> <given-names>D.-W.</given-names></name> <name><surname>Cheng</surname> <given-names>J.-H.</given-names></name></person-group> (<year>2016b</year>). <article-title>Recent advances in nondestructive analytical techniques for determining the total soluble solids in fruits: a review</article-title>. <source>Comprehensive Rev. Food Sci. Food Safety</source> <volume>15</volume>, <fpage>897</fpage>&#x02013;<lpage>911</lpage>. <pub-id pub-id-type="doi">10.1111/1541-4337.12217</pub-id><pub-id pub-id-type="pmid">33401799</pub-id></citation></ref>
<ref id="B21">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>L.</given-names></name> <name><surname>Huang</surname> <given-names>W.</given-names></name> <name><surname>Wang</surname> <given-names>Z.</given-names></name> <name><surname>Liu</surname> <given-names>S.</given-names></name> <name><surname>He</surname> <given-names>X.</given-names></name> <name><surname>Fan</surname> <given-names>S.</given-names></name></person-group> (<year>2022</year>). <article-title>Calibration transfer between developed portable vis/nir devices for detection of soluble solids contents in apple</article-title>. <source>Postharvest Biol. Technol</source>. 183, 111720. <pub-id pub-id-type="doi">10.1016/j.postharvbio.2021.111720</pub-id></citation>
</ref>
<ref id="B22">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>X.</given-names></name> <name><surname>Wei</surname> <given-names>Y.</given-names></name> <name><surname>Xu</surname> <given-names>J.</given-names></name> <name><surname>Feng</surname> <given-names>X.</given-names></name> <name><surname>Wu</surname> <given-names>F.</given-names></name> <name><surname>Zhou</surname> <given-names>R.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Ssc and ph for sweet assessment and maturity classification of harvested cherry fruit based on nir hyperspectral imaging technology</article-title>. <source>Postharvest Biol. Technol</source>. <volume>143</volume>, <fpage>112</fpage>&#x02013;<lpage>118</lpage>. <pub-id pub-id-type="doi">10.1016/j.postharvbio.2018.05.003</pub-id></citation>
</ref>
<ref id="B23">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ma</surname> <given-names>T.</given-names></name> <name><surname>Li</surname> <given-names>X.</given-names></name> <name><surname>Inagaki</surname> <given-names>T.</given-names></name> <name><surname>Yang</surname> <given-names>H.</given-names></name> <name><surname>Tsuchikawa</surname> <given-names>S.</given-names></name></person-group> (<year>2017</year>). <article-title>Noncontact evaluation of soluble solids content in apples by near-infrared hyperspectral imaging</article-title>. <source>J. Food Eng</source>. <volume>224</volume>, <fpage>53</fpage>&#x02013;<lpage>61</lpage>. <pub-id pub-id-type="doi">10.1016/j.jfoodeng.2017.12.028</pub-id></citation>
</ref>
<ref id="B24">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Ma</surname> <given-names>T.</given-names></name> <name><surname>Xia</surname> <given-names>Y.</given-names></name> <name><surname>Inagaki</surname> <given-names>T.</given-names></name> <name><surname>Tsuchikawa</surname> <given-names>S.</given-names></name></person-group> (<year>2021</year>). <article-title>Rapid and nondestructive evaluation of soluble solids content (ssc) and firmness in apple using vis-nir spatially resolved spectroscopy</article-title>. <source>Postharvest Biol. Technol</source>. 173, 111417. <pub-id pub-id-type="doi">10.1016/j.postharvbio.2020.111417</pub-id></citation>
</ref>
<ref id="B25">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mamouei</surname> <given-names>M.</given-names></name> <name><surname>Budidha</surname> <given-names>K.</given-names></name> <name><surname>Baishya</surname> <given-names>N.</given-names></name> <name><surname>Qassem</surname> <given-names>M.</given-names></name> <name><surname>Kyriacou</surname> <given-names>P.</given-names></name></person-group> (<year>2020</year>). <article-title>Comparison of wavelength selection methods for in-vitro estimation of lactate: a new unconstrained, genetic algorithm-based wavelength selection</article-title>. <source>Sci. Rep</source>. <volume>10</volume>, <fpage>1</fpage>&#x02013;<lpage>12</lpage>. <pub-id pub-id-type="doi">10.1038/s41598-020-73406-4</pub-id><pub-id pub-id-type="pmid">33037265</pub-id></citation></ref>
<ref id="B26">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Musingarabwi</surname> <given-names>D. M.</given-names></name> <name><surname>Nieuwoudt</surname> <given-names>H. H.</given-names></name> <name><surname>Young</surname> <given-names>P. R.</given-names></name> <name><surname>Ey&#x000E9;gh&#x000E8;-Bickong</surname> <given-names>H.</given-names></name> <name><surname>Vivier</surname> <given-names>M. A.</given-names></name></person-group> (<year>2016</year>). <article-title>A rapid qualitative and quantitative evaluation of grape berries at various stages of development using fourier-transform infrared spectroscopy and multivariate data analysis</article-title>. <source>Food Chem</source>. <volume>190</volume>, <fpage>253</fpage>&#x02013;<lpage>262</lpage>. <pub-id pub-id-type="doi">10.1016/j.foodchem.2015.05.080</pub-id><pub-id pub-id-type="pmid">26212968</pub-id></citation></ref>
<ref id="B27">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Nam</surname> <given-names>H.-A.</given-names></name> <name><surname>Ramakrishnan</surname> <given-names>S. R.</given-names></name> <name><surname>Kwon</surname> <given-names>J.-H.</given-names></name></person-group> (<year>2019</year>). <article-title>Effects of electron-beam irradiation on the quality characteristics of mandarin oranges (citrus unshiu (swingle) marcov) during storage</article-title>. <source>Food Chem</source>. <volume>286</volume>, <fpage>338</fpage>&#x02013;<lpage>345</lpage>. <pub-id pub-id-type="doi">10.1016/j.foodchem.2019.02.009</pub-id><pub-id pub-id-type="pmid">30827616</pub-id></citation></ref>
<ref id="B28">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Wei</surname> <given-names>X.</given-names></name> <name><surname>He</surname> <given-names>J.</given-names></name> <name><surname>Zheng</surname> <given-names>S.</given-names></name> <name><surname>Ye</surname> <given-names>D.</given-names></name></person-group> (<year>2020</year>). <article-title>Modeling for ssc and firmness detection of persimmon based on nir hyperspectral imaging by sample partitioning and variables selection</article-title>. <source>Infrared Phys. Technol</source>. 105, 103099. <pub-id pub-id-type="doi">10.1016/j.infrared.2019.103099</pub-id></citation>
</ref>
<ref id="B29">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Xia</surname> <given-names>Y.</given-names></name> <name><surname>Fan</surname> <given-names>S.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Tian</surname> <given-names>X.</given-names></name> <name><surname>Huang</surname> <given-names>W.</given-names></name> <name><surname>Chen</surname> <given-names>L.</given-names></name></person-group> (<year>2020</year>). <article-title>Optimization and comparison of models for prediction of soluble solids content in apple by online vis/nir transmission coupled with diameter correction method</article-title>. <source>Chemometr. Intell. Labo. Syst</source>. 201, 104017. <pub-id pub-id-type="doi">10.1016/j.chemolab.2020.104017</pub-id></citation>
</ref>
<ref id="B30">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xia</surname> <given-names>Z.</given-names></name> <name><surname>Wu</surname> <given-names>D.</given-names></name> <name><surname>Nie</surname> <given-names>P.</given-names></name> <name><surname>He</surname> <given-names>Y.</given-names></name></person-group> (<year>2016</year>). <article-title>Non-invasive measurement of soluble solid content and ph in kyoho grapes using a computer vision technique</article-title>. <source>Anal. Methods</source> <volume>8</volume>, <fpage>3242</fpage>&#x02013;<lpage>3248</lpage>. <pub-id pub-id-type="doi">10.1039/C5AY02694F</pub-id></citation>
</ref>
<ref id="B31">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>D.</given-names></name> <name><surname>Xu</surname> <given-names>Y.</given-names></name> <name><surname>Huang</surname> <given-names>W.</given-names></name> <name><surname>Tian</surname> <given-names>X.</given-names></name> <name><surname>Xia</surname> <given-names>Y.</given-names></name> <name><surname>Xu</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2019</year>). <article-title>Nondestructive measurement of soluble solids content in apple using near infrared hyperspectral imaging coupled with wavelength selection algorithm</article-title>. <source>Infrared Phys. Technol</source>. <volume>98</volume>, <fpage>297</fpage>&#x02013;<lpage>304</lpage>. <pub-id pub-id-type="doi">10.1016/j.infrared.2019.03.026</pub-id></citation>
</ref>
<ref id="B32">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>D.</given-names></name> <name><surname>Yang</surname> <given-names>Y.</given-names></name> <name><surname>Chen</surname> <given-names>G.</given-names></name> <name><surname>Tian</surname> <given-names>X.</given-names></name> <name><surname>Wang</surname> <given-names>Z.</given-names></name> <name><surname>Fan</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Nondestructive evaluation of soluble solids content in tomato with different stage by using vis/nir technology and multivariate algorithms</article-title>. <source>Spectrochim. Acta A</source> <volume>248</volume>, <fpage>119139</fpage>. <pub-id pub-id-type="doi">10.1016/j.saa.2020.119139</pub-id><pub-id pub-id-type="pmid">33214104</pub-id></citation></ref>
<ref id="B33">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Wang</surname> <given-names>J.</given-names></name> <name><surname>Ye</surname> <given-names>S.</given-names></name></person-group> (<year>2008</year>). <article-title>Predictions of acidity, soluble solids and firmness of pear using electronic nose technique</article-title>. <source>J. Food Eng</source>. <volume>86</volume>, <fpage>370</fpage>&#x02013;<lpage>378</lpage>. <pub-id pub-id-type="doi">10.1016/j.jfoodeng.2007.08.026</pub-id></citation>
</ref>
<ref id="B34">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>W.</given-names></name> <name><surname>Pan</surname> <given-names>L.</given-names></name> <name><surname>Zhao</surname> <given-names>X.</given-names></name> <name><surname>Tu</surname> <given-names>K.</given-names></name></person-group> (<year>2016</year>). <article-title>A study on soluble solids content assessment using electronic nose: persimmon fruit picked on different dates</article-title>. <source>Int. J. Food Propert</source>. <volume>19</volume>, <fpage>53</fpage>&#x02013;<lpage>62</lpage>. <pub-id pub-id-type="doi">10.1080/10942912.2014.940535</pub-id></citation>
</ref>
<ref id="B35">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhou</surname> <given-names>Q.</given-names></name> <name><surname>Huang</surname> <given-names>W.</given-names></name> <name><surname>Fan</surname> <given-names>S.</given-names></name> <name><surname>Zhao</surname> <given-names>F.</given-names></name> <name><surname>Liang</surname> <given-names>D.</given-names></name> <name><surname>Tian</surname> <given-names>X.</given-names></name></person-group> (<year>2020</year>). <article-title>Non-destructive discrimination of the variety of sweet maize seeds based on hyperspectral image coupled with wavelength selection algorithm</article-title>. <source>Infrared Phys. Technol</source>. 109, 103418. <pub-id pub-id-type="doi">10.1016/j.infrared.2020.103418</pub-id></citation>
</ref>
<ref id="B36">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zou</surname> <given-names>Z.</given-names></name> <name><surname>Xi</surname> <given-names>W.</given-names></name> <name><surname>Hu</surname> <given-names>Y.</given-names></name> <name><surname>Nie</surname> <given-names>C.</given-names></name> <name><surname>Zhou</surname> <given-names>Z.</given-names></name></person-group> (<year>2016</year>). <article-title>Antioxidant activity of citrus fruits</article-title>. <source>Food Chem</source>. <volume>196</volume>, <fpage>885</fpage>&#x02013;<lpage>896</lpage>. <pub-id pub-id-type="doi">10.1016/j.foodchem.2015.09.072</pub-id><pub-id pub-id-type="pmid">26593569</pub-id></citation></ref>
</ref-list>
</back>
</article>