<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Plant Sci.</journal-id>
<journal-title>Frontiers in Plant Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Plant Sci.</abbrev-journal-title>
<issn pub-type="epub">1664-462X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpls.2025.1473153</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Plant Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Evaluating sowing uniformity in hybrid rice using image processing and the OEW-YOLOv8n network</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Li</surname>
<given-names>Zehua</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2804814"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Pan</surname>
<given-names>Yihui</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2804194"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Ma</surname>
<given-names>Xu</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>*</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1332032"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Lin</surname>
<given-names>Yongjun</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2804803"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Wang</surname>
<given-names>Xicheng</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2805593"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Li</surname>
<given-names>Hongwei</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2552413"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>College of Mathematics and Informatics, South China Agricultural University</institution>, <addr-line>Guangzhou</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Key Laboratory of Smart Agricultural Technology in Tropical South China, Ministry of Agriculture and Rural Affairs</institution>, <addr-line>Guangzhou</addr-line>, <country>China</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>College of Engineering, South China Agricultural University</institution>, <addr-line>Guangzhou</addr-line>, <country>China</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>School of Mechanical Engineering, Guangxi University</institution>, <addr-line>Nanning</addr-line>, <country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>Edited by: Yunchao Tang, Dongguan University of Technology, China</p>
</fn>
<fn fn-type="edited-by">
<p>Reviewed by: Magdi A. A. Mousa, King Abdulaziz University, Saudi Arabia</p>
<p>Ting Yun, Nanjing Forestry University, China</p>
</fn>
<fn fn-type="corresp" id="fn001">
<p>*Correspondence: Xu Ma, <email xlink:href="mailto:maxu1959@scau.edu.cn">maxu1959@scau.edu.cn</email>
</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>03</day>
<month>02</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>16</volume>
<elocation-id>1473153</elocation-id>
<history>
<date date-type="received">
<day>12</day>
<month>08</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>06</day>
<month>01</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2025 Li, Pan, Ma, Lin, Wang and Li</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Li, Pan, Ma, Lin, Wang and Li</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Sowing uniformity is an important evaluation indicator of mechanical sowing quality. In order to achieve accurate evaluation of sowing uniformity in hybrid rice mechanical sowing, this study takes the seeds in a seedling tray of hybrid rice blanket-seedling nursing as the research object and proposes a method for evaluating sowing uniformity by combining image processing methods and the ODConv_C2f-ECA-WIoU-YOLOv8n (OEW-YOLOv8n) network. Firstly, image processing methods are used to segment seed image and obtain seed grids. Next, an improved model named OEW-YOLOv8n based on YOLOv8n is proposed to identify the number of seeds in a unit seed grid. The improved strategies include the following: (1) Replacing the Conv module in the Bottleneck of C2f modules with the Omni-Dimensional Dynamic Convolution (ODConv) module, where C2f modules are located at the connection between the Backbone and Neck. This improvement can enhance the feature extraction ability of the Backbone network, as the new modules can fully utilize the information of all dimensions of the convolutional kernel. (2) An Efficient Channel Attention (ECA) module is added to the Neck for improving the network&#x2019;s capability to extract deep semantic feature information of the detection target. (3) In the Bbox module of the prediction head, the Complete Intersection over Union (CIoU) loss function is replaced by the Weighted Intersection over Union version 3 (WIoUv3) loss function to improve the convergence speed of the bounding box loss function and reduce the convergence value of the loss function. The results show that the mean average precision (mAP) of the OEW-YOLOv8n network reaches 98.6%. Compared to the original model, the mAP improved by 2.5%. Compared to the advanced object detection algorithms such as Faster-RCNN, SSD, YOLOv4, YOLOv5s YOLOv7-tiny, and YOLOv10s, the mAP of the new network increased by 5.2%, 7.8%, 4.9%, 2.8% 2.9%, and 3.3%, respectively. Finally, the actual evaluation experiment showed that the test error is from &#x2212;2.43% to 2.92%, indicating that the improved network demonstrates excellent estimation accuracy. The research results can provide support for the mechanized sowing quality detection of hybrid rice and the intelligent research of rice seeder.</p>
</abstract>
<kwd-group>
<kwd>mechanical sowing</kwd>
<kwd>uniformity evaluation</kwd>
<kwd>deep learning</kwd>
<kwd>object detection</kwd>
<kwd>rice seeder</kwd>
</kwd-group>
<contract-num rid="cn001">52175226, 2023B0202130001</contract-num>
<contract-sponsor id="cn001">South China Agricultural University<named-content content-type="fundref-id">10.13039/501100012601</named-content>
</contract-sponsor>
<counts>
<fig-count count="13"/>
<table-count count="5"/>
<equation-count count="17"/>
<ref-count count="30"/>
<page-count count="17"/>
<word-count count="10125"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-in-acceptance</meta-name>
<meta-value>Sustainable and Intelligent Phytoprotection</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1" sec-type="intro">
<label>1</label>
<title>Introduction</title>
<p>Hybrid rice is an important food crop, accounting for over 50% of China&#x2019;s rice planting area. It has been planted in over 70 countries worldwide, with a cumulative planting area exceeding 600 million ha, making significant contributions to global food security (<xref ref-type="bibr" rid="B15">Ma and Yuan, 2015</xref>; <xref ref-type="bibr" rid="B25">Yuan, 2018</xref>). However, the quality of mechanical sowing of hybrid rice is one of the important factors restricting the development of hybrid rice production (<xref ref-type="bibr" rid="B4">Chen et&#xa0;al., 2015</xref>; <xref ref-type="bibr" rid="B13">Li et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B14">Ma et&#xa0;al., 2023</xref>) where sowing uniformity is an important evaluation indicator of the quality of mechanized sowing (<xref ref-type="bibr" rid="B1">Alekseev et&#xa0;al., 2018</xref>). Currently, the evaluation of the sowing uniformity in hybrid rice mainly relies on manual visual inspection, which is highly subjective. In order to achieve an objective evaluation of sowing uniformity, provide timely feedback on the sowing quality information to the performance control system of the seeders, and effectively improve sowing quality, this paper combines image processing methods and deep learning algorithms to study a method for evaluating the sowing uniformity of hybrid rice blanket-seedling nursing.</p>
<p>In the mechanized seedling nursing of hybrid rice, according to the different types of seedling trays, seedling nursing methods are divided into pot-seedling nursing and blanket-seedling nursing, where pot trays and blanket trays are used, respectively. The pot tray is a type of plastic tray. It is usually composed of 406 (29 rows&#xd7;14 columns) or 448 (32 rows&#xd7;14 columns) pots arranged uniformly. The blanket tray is also a type of plastic tray, which is essentially an uncovered cuboid with a typical inner cavity size of 25 mm &#xd7; 280 mm &#xd7; 580 mm. When a pot tray is used for sowing, the seeds in the tray are separated into relatively separate seed groups by pots, while when a blanket tray is used for sowing, the seeds in the tray are relatively uniformly arranged, but there is still a certain degree of randomness in the direction and position of the seeds. This indicates that compared to blanket-seedling nursing, the target area of sowing quality inspection for pot-seedling nursing is relatively easy to be determined. Therefore, both domestically and internationally, research on the sowing quality detection of hybrid rice mechanized sowing has mainly focused on pot-seedling nursing (<xref ref-type="bibr" rid="B17">Qi et&#xa0;al., 2009</xref>; <xref ref-type="bibr" rid="B30">Zhou et&#xa0;al., 2012</xref>; <xref ref-type="bibr" rid="B28">Zhao et&#xa0;al., 2014</xref>; <xref ref-type="bibr" rid="B20">Wang et&#xa0;al., 2015</xref>; <xref ref-type="bibr" rid="B2">Chen et&#xa0;al., 2017</xref>; <xref ref-type="bibr" rid="B21">Wang et&#xa0;al., 2018</xref>), in which the combination of image processing and machine vision technology is mainly used. There are relatively fewer reports on the sowing quality detection of blanket-seedling nursing, and only articles published by <xref ref-type="bibr" rid="B18">Tan et&#xa0;al. (2014)</xref> and <xref ref-type="bibr" rid="B6">Dong et&#xa0;al. (2020)</xref> have been found. Among them, <xref ref-type="bibr" rid="B18">Tan et&#xa0;al. (2014)</xref> established a BP neural network model for detecting the sowing quantity of super hybrid rice pot-seedling nursing by extracting shape features such as area, perimeter, and shape factor of the seed-connected region. The average accuracy of detection has reached 94.4%. <xref ref-type="bibr" rid="B6">Dong et&#xa0;al. (2020)</xref> designed a sowing quantity detection and control device for hybrid rice based on embedded machine vision, where the algorithm for detecting the sowing quantity is similar to the algorithm in <xref ref-type="bibr" rid="B18">Tan et&#xa0;al. (2014)</xref>.</p>
<p>In recent years, deep learning has demonstrated excellent performance in complex scene object detection and has been widely applied in agricultural production (<xref ref-type="bibr" rid="B8">Johansen et&#xa0;al., 2019a</xref>, <xref ref-type="bibr" rid="B9">b</xref>, <xref ref-type="bibr" rid="B10">2020</xref>; <xref ref-type="bibr" rid="B7">Jiang et&#xa0;al., 2022</xref>; <xref ref-type="bibr" rid="B26">Yun et&#xa0;al., 2024</xref>). Although no studies using deep learning methods to investigate the sowing uniformity of hybrid rice blanket-seedling nursing have been found, deep learning methods have been applied for uniformity detection in other crop production, including tasks such as corn seedling emergence uniformity detection. For instance, <xref ref-type="bibr" rid="B16">Nee et&#xa0;al. (2022)</xref> utilized UAV imagery and deep learning model ResNet18 to estimate and map corn emergence uniformity. Corn emergence uniformity was quantified with plant density, plant spacing standard deviation, and mean days to imaging after emergence.</p>
<p>You Only Look Once (YOLO) is a typical representative of one-stage networks in deep learning, characterized by high detection accuracy and fast speed, and has been widely used in non-contact object detection. For example, <xref ref-type="bibr" rid="B23">Wang et&#xa0;al. (2023)</xref> proposed an improved YOLOv5 model named CM-YOLOv5s-CSVPoVnet model for yellow peach detection by adding model clipping, Omni-Dimensional Dynamic Convolution (ODConv), Global Context Networks, and ELAN methods. The results showed that the mean average precision (mAP) reached 96%, and the model size was 3.54 MB. <xref ref-type="bibr" rid="B27">Zhang et&#xa0;al. (2023)</xref> enhanced the YOLOv5s model by incorporating the Efficient Channel Attention (ECA)-Net attention mechanism module and Adaptively Spatial Feature Fusion (ASFF) into the feature pyramid structure of YOLO, thus creating the YOLOv5-ECA-ASFF model. This model was utilized for detecting wheat scab fungus spores, and the results demonstrated that the average recognition accuracy reached 98.57%. The YOLO series has algorithms such as YOLOv1 to YOLOv10, with the characteristic that newer versions are more advanced. YOLO has demonstrated excellent detection performance in the above studies, being able to quickly count the number of targets, which is very enlightening for detecting the sowing uniformity of hybrid rice.</p>
<p>In this paper, according to the characteristics of sowing quality detection of hybrid rice blanket-seedling nursing, an improved algorithm named ODConv_C2f-ECA-WIoU-YOLOv8n (OEW-YOLOv8n) network is proposed based on the YOLOv8n algorithm for identifying the number of seeds in a unit seed grid. Then, combining image processing methods with the OEW-YOLOv8n network, a method for evaluating sowing uniformity of hybrid rice blanket-seedling nursing is promoted. The sowing uniformity of detection in seed image has been implemented on computer devices. The relevant results can provide support for detecting the sowing quality of hybrid rice.</p>
</sec>
<sec id="s2" sec-type="materials|methods">
<label>2</label>
<title>Materials and methods</title>
<sec id="s2_1">
<label>2.1</label>
<title>Mechanical sowing and image acquisition</title>
<p>The mechanical sowing experiment was conducted on 1 July 2023 at the Soil Trough Laboratory of College of Engineering, South China Agricultural University, Guangdong Province (latitude 23&#xb0;16&#x2032;, longitude 113&#xb0;35&#x2032;). The time for mechanical sowing is from 9:00 a.m. to 10:00 a.m., and the time for image acquisition is from 10:00 a.m. to 12:00 p.m. The used seeder is the 2ZSB-500 intelligent rice tray seedling precise seeding production line developed by South China Agricultural University. The experimental variety is hybrid rice taifengyou 208. Before mechanical sowing, the seeds are carefully selected and germinated. The seedling trays are standard hard blanket-seedling trays with specifications of 25 mm &#xd7; 280 mm &#xd7; 580 mm. The mechanical sowing process mainly includes laying tray, laying bottom soil, and sowing. After sowing, the seeds are not covered by topsoil for capturing seed images. The seedling trays with seeds are placed on a flat ground. Under natural light conditions (during the experiment, the strength of illumination was between 10,000 and 30,000 lux), the seed images of the entire seedling tray were acquired by a smartphone fixed on a simple camera perch. The smartphone is an Apple iPhone 12, with a resolution of 2,268 &#xd7; 4,032 pixels. In order to improve the generalization ability of the model by increasing the diversity of target data, the sowing methods include broadcast sowing, ditching strip sowing, and no-ditching strip sowing, and the sowing density consists of three levels: 50 g/tray, 60 g/tray, and 70 g/tray. Thus, there are nine kinds of sowing modes. For each sowing mode, 5 trays were sown, resulting in a total of 45 trays having been sown. For each tray, five photographs with different shooting effects were selected. Finally, we obtained 225 valid photographs. Among them, 18 photographs were used as training samples, consisting of 2 photographs of trays from 9 sowing modes. The remaining 207 photographs were used as application samples to evaluate sowing uniformity. The representative seed images and mechanical sowing process are shown in <xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref>.</p>
<fig id="f1" position="float">
<label>Figure&#xa0;1</label>
<caption>
<p>Representative seed images and mechanical sowing process. <bold>(A)</bold> demonstrates the mechanical sowing process, where a-i sequentially represents the seedling tray, automatic laying tray, laying bottom soil, precision sowing, laying topsoil, cleaning topsoil, automatic stacking tray, and precision sowing. In <bold>(B)</bold>, the seed images are shown, where i, ii, and iii represent broadcast sowing with 50 g/tray, ditching strip sowing with 60&#xa0;g/tray, and no-ditching strip sowing with 70 g/tray, respectively.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g001.tif"/>
</fig>
</sec>
<sec id="s2_2">
<label>2.2</label>
<title>The main process of evaluating sowing uniformity</title>
<p>The evaluation of sowing uniformity mainly includes three parts: image preprocessing, model training, and uniformity evaluation, as shown in <xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>. Image preprocessing is shown in <xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref>.</p>
<fig id="f2" position="float">
<label>Figure&#xa0;2</label>
<caption>
<p>The main process of evaluating sowing uniformity. (i) image preprocessing (ii) labeling (iii) data augmentation (iv) model training and testing (v) detect unit seed grid images (vi) evaluate uniformity Modification of problem.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g002.tif"/>
</fig>
<fig id="f3" position="float">
<label>Figure&#xa0;3</label>
<caption>
<p>Image preprocessing.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g003.tif"/>
</fig>
<p>The image preprocessing process is as follows: (1) Inputting seed images and obtaining binary images by using binarization. (2) Using Canny operator to perform edge detection on binary images and obtain edge contour images. (3) Using contour filtering to obtain the edge frame of seedling tray images. (4) Correcting the graphics through affine transformation to obtain seedling tray images, which includes detection area and the borders of seedling tray. (5) Locating the detection area through the vertical projection method (see Section 2.3.1). (6) Setting the number of grids, use the Sliding Window method to segment the detection area and obtain unit seed grid images (see Section 2.3.2).</p>
<p>The model training mainly includes (1) dividing the unit seed grid images into training samples and application samples, (2) using the LabelImg software to label the seeds in training samples, (3)&#xa0;using image augmentation methods to expand the training set, and (4) inputting the dataset into the OEW-YOLOv8n model for training and testing.</p>
<p>The uniformity evaluation part mainly includes the following: The seed images of the application samples are treated undergoing the image preprocessing process, and the seed grid images of the application samples will be obtained. Then, input the seed grid images of the application samples into the OEW-YOLOv8n model for detection, and the number of seeds per grid will be obtained. Finally, calculate the qualification rates of grids of one to three seeds per grid to evaluate the sowing uniformity.</p>
</sec>
<sec id="s2_3">
<label>2.3</label>
<title>Methods for detection area localization and seed grid segmentation based on image processing</title>
<sec id="s2_3_1">
<label>2.3.1</label>
<title>Detection area localization algorithm</title>
<p>Because the seedling tray image includes some non-detection areas such as the border of the seedling tray, it is necessary to locate the detection area. The vertical projection method (<xref ref-type="bibr" rid="B3">Chen et&#xa0;al., 2021</xref>) is used to solve this problem in this paper.</p>
<p>The vertical projection method is as follows: In the binary seedling tray image (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4</bold>
</xref>), white oval dots represent seeds with a grayscale value of 0, while black areas represent non-seed parts such as seedling soil or the border of the seedling tray, with a grayscale value of 1. The cumulative grayscale values for each row and each column are computed in pixel units, following the methods described in <xref ref-type="disp-formula" rid="eq1">Equations 1</xref> and <xref ref-type="disp-formula" rid="eq2">2</xref>.</p>
<fig id="f4" position="float">
<label>Figure&#xa0;4</label>
<caption>
<p>Binary seedling tray image.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g004.tif"/>
</fig>
<disp-formula id="eq1">
<label>(1)</label>
<mml:math display="block" id="M1">
<mml:mrow>
<mml:mi>w</mml:mi>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo stretchy="false">)</mml:mo>
<mml:mo>=</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>0</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>H</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:munderover>
<mml:mrow>
<mml:mi>p</mml:mi>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>,</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>,</mml:mo>
<mml:mi>W</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq2">
<label>(2)</label>
<mml:math display="block" id="M2">
<mml:mrow>
<mml:mi>h</mml:mi>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>j</mml:mi>
<mml:mo stretchy="false">)</mml:mo>
<mml:mo>=</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>0</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>W</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:munderover>
<mml:mrow>
<mml:mi>p</mml:mi>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>,</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>,</mml:mo>
<mml:mi>H</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</disp-formula>
<p>where <italic>W</italic> and <italic>H</italic> represent the number of rows and columns, respectively. In this paper, <italic>W</italic> = 4,032 and <italic>H</italic> = 2,268. <italic>w</italic>(<italic>i</italic>) represents the cumulative grayscale value of the <italic>i</italic>th row, <italic>h</italic>(<italic>j</italic>) represents the cumulative grayscale value of the <italic>j</italic>th column, and <italic>p</italic>(<italic>i</italic>,<italic>j</italic>) represents the grayscale value of the corresponding pixel at the <italic>i</italic>th row and <italic>j</italic>th column.</p>
<p>Based on the grayscale values of the binary image, a vertical projection curve (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5A</bold>
</xref>) was plotted with the number of horizontal pixels as abscissa and the cumulative grayscale value of each column as ordinate. Similarly, a horizontal projection curve (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5B</bold>
</xref>) was plotted with the number of vertical pixels as abscissa and the cumulative grayscale value of each row as ordinate. As shown in <xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5</bold>
</xref>, there are four abrupt segments on both sides of the images of the vertical projection curve and the horizontal projection, namely, A<sub>1</sub>B<sub>1</sub>, C<sub>1</sub>D<sub>1</sub>, or A<sub>2</sub>B<sub>2</sub>, C<sub>2</sub>D<sub>2</sub>. In segments A<sub>1</sub>B<sub>1</sub> and A<sub>2</sub>B<sub>2</sub>, the cumulative grayscale value suddenly increases from 0 to a larger value, owing to the appearance of the border of the seedling tray in the image, in which A<sub>1</sub> and A<sub>2</sub> represent the outer side of the seedling tray border, while B<sub>1</sub> and B<sub>2</sub> represent the inner side of the seedling tray border. Conversely, in segments C<sub>1</sub>D<sub>1</sub> and C<sub>2</sub>D<sub>2</sub>, the cumulative grayscale value suddenly drops from a larger value to 0. The reason for this is that the border of the seedling tray in the image disappears. At this point, C<sub>1</sub> and C<sub>2</sub> represent the inside of the seedling tray border, while D<sub>1</sub> and D<sub>2</sub> represent the outside of the seedling tray border. This indicates that A<sub>1</sub>B<sub>1</sub>, C<sub>1</sub>D<sub>1</sub>, A<sub>2</sub>B<sub>2</sub>, and C<sub>2</sub>D<sub>2</sub> respectively represent the positions of the four borders of the seedling tray. Taking the positions of the pixel numbers corresponding to B<sub>1</sub>, C<sub>1</sub>, B<sub>2</sub>, and C<sub>2</sub> as the positions of the four edge lines of the detection area, then the detection area is located.</p>
<fig id="f5" position="float">
<label>Figure&#xa0;5</label>
<caption>
<p>Projection curve. <bold>(A)</bold> vertical projection curve <bold>(B)</bold> horizontal projection curve.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g005.tif"/>
</fig>
</sec>
<sec id="s2_3_2">
<label>2.3.2</label>
<title>Seed grid segmentation method</title>
<p>In general, in order to distinguish the description of pots in a pot tray, it is customary to refer to the unit area of blanket tray as the seed grid. A seed grid in a blanket tray is equivalent to one pot in the pot tray. In order to evaluate the sowing uniformity of blanket-seedling nursing, it is necessary to manually divide seed grids and mark them in seed images.</p>
<p>Without considering the deformation of seedling blocks, one seed grid is equal to the unit seedling-taken area during mechanical transplanting. A unit seedling-taken area of mechanical transplanting = horizontal seedling cutting quantity &#xd7; longitudinal seedling cutting quantity. Therefore, the key to dividing seed grids is to determine the horizontal seedling cutting quantity and the longitudinal seedling cutting quantity. In theory, the number of seeds in the seed grid is the basis for estimating the number of seedlings per hill in the field during mechanical transplanting. Practical experience shows that the actual germination rate of hybrid rice seeds is from 75% to 85% (<xref ref-type="bibr" rid="B12">Li et&#xa0;al., 2019</xref>). To achieve high-yield cultivation of hybrid rice, theoretically, the number of seeds in one seed grid is relatively suitable at 2.7&#x2013;2.9. At this time, the average number of seedlings per hill is 2&#x2013;2.5. Therefore, the method for dividing seed grids in seed images is as follows:</p>
<p>In hybrid rice blanket-seedling nursing with strip sowing, an 18-row seeder is usually used for sowing. After sowing, the seeds in the seedling tray appear as 18 seed strips. Therefore, the seed image is evenly divided into 18 parts horizontally. At this time, the number of cutting times of horizontal seedling by seedling claw of the rice transplanter is 18, and the horizontal seedling cutting quantity is 15.6 mm (208 mm &#xf7; 18 = 15.6 mm). The number of cutting times of longitudinal seedling by seedling claw of the rice transplanter is determined by the longitudinal seedling cutting quantity. The longitudinal seedling cutting quantity is set by the longitudinal seedling feeding mechanism of the transplanter. Usually, the adjustment range for the longitudinal seedling feeding mechanism is from 8 to 18 mm. There are 11 gears available, with a 1-mm interval between each gear.</p>
<p>In order to achieve the target of precise transplanting of one to three seedlings per hill, the number of seed grids per tray is related to the sowing density. The higher the sowing density, the more the seed grids there will be. According to the thousand-grain weight and sowing density of hybrid rice seeds, the number and the division method of seed grids under different sowing densities are calculated as shown in <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref>.</p>
<table-wrap id="T1" position="float">
<label>Table&#xa0;1</label>
<caption>
<p>Number of seed grids in a seedling tray and division method under different sowing densities.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="left">Sowing density/(g&#xb7;tray&#x2212;1)</th>
<th valign="middle" align="left">50</th>
<th valign="middle" align="left">60</th>
<th valign="middle" align="left">70</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="left">The number of cutting times of horizontal seedling by seedling claw of the rice transplanter/times</td>
<td valign="middle" align="left">18</td>
<td valign="middle" align="left">18</td>
<td valign="middle" align="left">18</td>
</tr>
<tr>
<td valign="middle" align="left">Longitudinal seedling cutting quantity/mm</td>
<td valign="middle" align="left">14</td>
<td valign="middle" align="left">12</td>
<td valign="middle" align="left">10</td>
</tr>
<tr>
<td valign="middle" align="left">The number of cutting times of longitudinal seedling by seedling claw of the rice transplanter/times</td>
<td valign="middle" align="left">41</td>
<td valign="middle" align="left">48</td>
<td valign="middle" align="left">58</td>
</tr>
<tr>
<td valign="middle" align="left">Theoretically, the number of seed grids per tray/grid</td>
<td valign="middle" align="left">738</td>
<td valign="middle" align="left">864</td>
<td valign="middle" align="left">1,044</td>
</tr>
<tr>
<td valign="middle" align="left">Theoretically, the number of seeds per tray/grain</td>
<td valign="middle" align="left">2,020</td>
<td valign="middle" align="left">2,424</td>
<td valign="middle" align="left">2,828</td>
</tr>
<tr>
<td valign="middle" align="left">Theoretically, the number of seeds per grid/grain</td>
<td valign="middle" align="left">2.74</td>
<td valign="middle" align="left">2.81</td>
<td valign="middle" align="left">2.71</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>In <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref>, the thousand-grain weight of hybrid rice taifengyou 208 is 24.75 g. When the sowing density is 50 g/tray, theoretically, there are approximately 2,020 seeds per tray. If the number of seed grids per tray is set to 738 grids (18 rows &#xd7; 41 columns = 738 grids), then the average number of seeds per grid is 2.74, which meets the agronomic requirements. Similarly, when the sowing density is 60 g/tray and 70 g/tray, the seed grids per tray is 864 grids (18 rows &#xd7; 48 columns) and 1,044 grids (18 rows &#xd7; 58 columns), respectively. Theoretically, the number of seeds per grid is approximately 2.81 and 2.71, respectively.</p>
</sec>
</sec>
<sec id="s2_4">
<label>2.4</label>
<title>Dataset creation</title>
<p>After localizing the detection areas of the training samples for the model, we get 18 seed images only containing detection areas. According to the seed grid division methods of seedling tray under different sowing densities in <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref>, using the Sliding Window method to divide the seed images, we obtained 15,876 unit seed grid images [(738 + 864 + 1,044) &#xd7; 3 &#xd7; 2 = 15,876].</p>
<p>This paper uses the LabelImg tool to label seeds in seed grid images, naming the labels &#x201c;seed&#x201d; and generating corresponding txt label files. Expand the training sample dataset using data augmentation methods, which mainly include cropping, adding Gaussian noise, horizontal flipping, vertical flipping, color transformation, contrast transformation, and randomly scaling the width and height of seed grid images within a reasonable range. The expanded training dataset includes 23,814 images with a total of 82,196 seed labels. The expanded images and their corresponding text label files are randomly divided into a training set and a testing set in a 7:2:1 ratio (as shown in <xref ref-type="table" rid="T2">
<bold>Table&#xa0;2</bold>
</xref>).</p>
<table-wrap id="T2" position="float">
<label>Table&#xa0;2</label>
<caption>
<p>Distribution of seed grid images for training sample.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="center">Categories</th>
<th valign="middle" align="center">Number of images</th>
<th valign="middle" align="center">Number of labels</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="center">Training set</td>
<td valign="middle" align="center">16,669</td>
<td valign="middle" align="center">57,535</td>
</tr>
<tr>
<td valign="middle" align="center">Validation set</td>
<td valign="middle" align="center">4,763</td>
<td valign="middle" align="center">16,441</td>
</tr>
<tr>
<td valign="middle" align="center">Test set</td>
<td valign="middle" align="center">2,382</td>
<td valign="middle" align="center">8,220</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s2_5">
<label>2.5</label>
<title>OEW-YOLOv8n model</title>
<p>YOLOv8 is an updated version of YOLOv5 developed by Ultralytics company. Compared to the previous-generation YOLOv5, the main improvements in YOLOv8 include the following: (1) In the Backbone network, the C3 module has been replaced by the C2f module, achieving further lightweight; (2) in the Neck, the Path Aggregation Network&#x2013;Feature Pyramid Network (PA-FPN) concept has been adopted, removing the convolution operations during the upsampling process; (3) the Decoupled-Head structure is adopted in the Head, separating the classification and detection tasks, further reducing the complexity of the model; and (4) in the loss function, binary cross-entropy loss (BCE Loss) is used for classification, and distribution focal loss (DF Loss) along with Complete Intersection over Union (CIoU) Loss is used for regression. These improvements allow YOLOv8 to retain the advantages of the YOLOv5 network structure while making more refined adjustments and optimizations, enhancing the model&#x2019;s performance in various scenarios. YOLOv8 is a family of models, ranging from smaller to larger versions, including YOLOv8n, YOLOv8s, YOLOv8m, YOLOv8l, and YOLOv8x. The key in evaluating sowing uniformity in hybrid rice blanket-seedling nursing is identifying the number of seeds in a unit seed grid. To ensure the model is lightweight, this study constructs the model based on YOLOv8n, which has the minimum number of model parameters.</p>
<p>The YOLOv8n network mainly consists of four parts: Input, Backbone, Neck, and Head. The Input terminal enhances the unit seed grid image through Mosaic augmentation before feeding it into the network. The Backbone network replaces the C3 module with the C2f module and uses the CSPDarkNet53 network to extract features from top to bottom. Additionally, it includes five Conv modules and one SPPF module. The C2f module contains two Conv modules, one split module, <italic>n</italic> Bottleneck modules, and one Concat operation. The C2f module captures rich gradient information from the feature map. The SPPF module uses convolutional kernels of different sizes to pool the feature map, obtaining multi-scale features. It then utilizes the Concat operation to achieve multi-scale feature fusion. The Neck network mainly consists of the PAN. PAN consists of the Feature Pyramid Network (FPN) and a bottom-up PAN. FPN passes rich semantic information from the top to the bottom, enhancing the semantic information of shallow feature maps, while the bottom-up PAN propagates strong localization feature information from the bottom up. The dual-pyramid structure of the Neck network further enhances the representation capability of multi-scale features. The prediction head outputs the classification and coordinate information of detected targets in three different size branches. The classification loss function uses BCE Loss, and the object detection regression loss functions use DF Loss and CIoU Loss.</p>
<p>The aim of this study is to identify the number of seeds in a unit seed grid. Hybrid rice seed images have the following characteristics: (1) Each image contains a large number of seeds, with individual seeds occupying relatively few pixels, which falls under the category of small object detection. (2) Because of seed grid division, some seeds are artificially cut, so some seeds are incomplete in a unit seed grid, resulting in diversity of the targets. (3) Some seeds exhibit crossing, overlapping, or occlusion in the images. (4) The background of seed images is relatively complex for the nursery soil and consists of different components with different colors and shapes. Based on the above characteristics, this paper proposes the following improvements to YOLOv8n: (1) Adjust the C2f module connected to Neck in the Backbone structure and replace the Conv module in the Bottleneck structure with the ODConv module (<xref ref-type="bibr" rid="B11">Li et&#xa0;al., 2022</xref>), so that this module fully utilizes the information of all dimensions of the convolutional kernels, enhancing the feature extraction capability of the model&#x2019;s Backbone network and improving the accuracy of small object detection. (2) Add an ECA attention mechanism module (<xref ref-type="bibr" rid="B22">Wang et&#xa0;al., 2020</xref>) to the Neck structure to enhance the model&#x2019;s ability to extract deep semantic feature information of detection targets and improve the accuracy of detecting incomplete targets. (3) In the Bbox-Loss module of Head, the WIoUv3 loss function (<xref ref-type="bibr" rid="B19">Tong et&#xa0;al., 2023</xref>) is used instead of the CIoU loss function (<xref ref-type="bibr" rid="B29">Zheng et&#xa0;al., 2020</xref>) to improve the convergence speed of the bounding box loss function and reduce the convergence value of the loss function. The network structure is shown in <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>, and this algorithm is named ODConv_C2f-ECA-WIoU-YOLOv8n, abbreviated as the OEW-YOLOv8n model.</p>
<fig id="f6" position="float">
<label>Figure&#xa0;6</label>
<caption>
<p>The architecture of OEW-YOLOv8n. The improved modules are highlighted in red, green, or blue colors, while the white boxes represent the original network structure modules. The improved Bottleneck module is named ODC_Bottleneck, and the improved C2f module is named ODC_C2f.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g006.tif"/>
</fig>
<sec id="s2_5_1">
<label>2.5.1</label>
<title>ODConv dynamic convolution module</title>
<p>The full name of ODConv is Omni Dimensional Dynamic Convolution, also known as Full Dimensional Dynamic Convolution. This method was first introduced in the article by <xref ref-type="bibr" rid="B11">Li et&#xa0;al. (2022)</xref>. The motivation for introducing this module in this article is as follows:</p>
<p>In YOLOv8n, the standard Conv is adopted by the C2f module, which essentially uses the same static convolution kernel to learn from the input samples in each convolution layer. At this point, the extraction of sample feature information is insufficient. To address this, researchers have studied dynamic convolutions, such as DyConv (<xref ref-type="bibr" rid="B5">Chen et&#xa0;al., 2020</xref>) and CondConv (<xref ref-type="bibr" rid="B24">Yang et&#xa0;al., 2019</xref>). The essence of dynamic convolution is to introduce an attention mechanism that learns a linear combination of <italic>n</italic> convolutional kernels with attention as weights. For example, the DyConv is defined as:</p>
<disp-formula id="eq3">
<label>(3)</label>
<mml:math display="block" id="M3">
<mml:mrow>
<mml:mi>y</mml:mi>
<mml:mo>=</mml:mo>
<mml:mo stretchy="false">(</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>w</mml:mi>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:msub>
<mml:mi>W</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>+</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>+</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>w</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
<mml:msub>
<mml:mi>W</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo stretchy="false">)</mml:mo>
<mml:mo>*</mml:mo>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:math>
</disp-formula>
<p>where <inline-formula>
<mml:math display="inline" id="im1">
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:msup>
<mml:mtext>R</mml:mtext>
<mml:mrow>
<mml:mi>h</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>w</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:msub>
<mml:mi>c</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula>
<mml:math display="inline" id="im2">
<mml:mrow>
<mml:mi>y</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:msup>
<mml:mtext>R</mml:mtext>
<mml:mrow>
<mml:mi>h</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>w</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:msub>
<mml:mi>c</mml:mi>
<mml:mrow>
<mml:mi>o</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> represent the input feature and output feature respectively. <italic>h</italic> and <italic>w</italic> represent the height and width of the channels, respectively. <italic>c<sub>in</sub>
</italic> and <italic>c<sub>out</sub>
</italic> represent the number of input channels and output channels, respectively. <italic>W<sub>i</sub>
</italic> represents the <italic>i</italic>th convolutional kernel. <italic>&#x3b1;<sub>wi</sub>
</italic> is an attention scalar that weights <italic>W<sub>i</sub>
</italic>, where <italic>i</italic> = 1,&#x2026;, <italic>n</italic>, and * represents the convolution operation. Research has shown that dynamic convolution can effectively improve the accuracy of lightweight convolutional neural networks while ensuring efficient inference (<xref ref-type="bibr" rid="B24">Yang et&#xa0;al., 2019</xref>; <xref ref-type="bibr" rid="B5">Chen et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B11">Li et&#xa0;al., 2022</xref>).</p>
<p>However, dynamic convolution only improves the convolution operation in terms of the number of convolutional kernels, neglecting the spatial kernel size dimension, input channel dimension, and output channel dimension. This reduces the ability of the convolution operation to extract information about sample characteristics. To address this limitation, <xref ref-type="bibr" rid="B11">Li et&#xa0;al. (2022)</xref> proposed ODConv based on multi-dimensional attention mechanisms and parallel strategies, defined as:</p>
<disp-formula id="eq4">
<label>(4)</label>
<mml:math display="block" id="M4">
<mml:mrow>
<mml:mtable>
<mml:mtr columnalign="left">
<mml:mtd columnalign="left">
<mml:mi>y</mml:mi>
<mml:mo>=</mml:mo>
<mml:mo stretchy="false">(</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>w</mml:mi>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2a00;</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>f</mml:mi>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2a00;</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2a00;</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>s</mml:mi>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2a00;</mml:mo>
<mml:msub>
<mml:mi>W</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>+</mml:mo>
<mml:mo>&#x2026;</mml:mo>
</mml:mtd>
</mml:mtr>
<mml:mtr columnalign="left">
<mml:mtd columnalign="left">
<mml:mo>+</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>w</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2a00;</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>f</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2a00;</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2a00;</mml:mo>
<mml:msub>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>s</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2a00;</mml:mo>
<mml:msub>
<mml:mi>W</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo stretchy="false">)</mml:mo>
<mml:mo>*</mml:mo>
<mml:mi>x</mml:mi>
</mml:mtd>
</mml:mtr>
</mml:mtable>
</mml:mrow>
</mml:math>
</disp-formula>
<p>
<italic>x</italic> represents the input features, while <italic>y</italic> represents the output features. <italic>W<sub>i</sub>
</italic> represents the <italic>i</italic>th convolutional kernel. <italic>&#x3b1;<sub>wi</sub>
</italic>, <italic>&#x3b1;<sub>fi</sub>
</italic>, <italic>&#x3b1;<sub>ci</sub>
</italic>, and <italic>&#x3b1;<sub>si</sub>
</italic> are four attention scalars, respectively, representing the weights of the convolutional kernel <italic>W<sub>i</sub>
</italic> (<italic>i</italic> = 1, &#x2026;, <italic>n</italic>) in the dimensions of kernel quantity, output channels, input channels, and spatial kernel size. The structure is shown in <xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>, where &#x2a00; represents the multiplication operator across different dimensions in the kernel space, and * represents the convolution operation.</p>
<fig id="f7" position="float">
<label>Figure&#xa0;7</label>
<caption>
<p>ODConv structure.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g007.tif"/>
</fig>
<p>In brief, in ODConv, <italic>&#x3b1;<sub>si</sub>
</italic> assigns different attention to the filters (convolution parameters) at each spatial location (<italic>k</italic> &#xd7; <italic>k</italic> in total, where <italic>k</italic> is the kernel size). <italic>&#x3b1;<sub>ci</sub>
</italic> assigns different attention to each input channel (<italic>C<sub>in</sub>
</italic>). <italic>&#x3b1;<sub>fi</sub>
</italic> assigns different attention to each output channel (<italic>C<sub>out</sub>
</italic>). <italic>&#x3b1;<sub>wi</sub>
</italic> assigns different attention to each convolutional kernel. These four types of attention are complementary and are multiplied with the convolutional kernel <italic>W<sub>i</sub>
</italic> in a sequential order of position, channel, filter, and kernel. This provides performance guarantees for capturing rich temporal and spatial cues. In summary, compared to mainstream dynamic convolutions like DyConv and CondConv, ODConv enhances feature learning capabilities and improves the accuracy of convolutional neural networks by obtaining complementary attention for the convolutional kernels along all dimensions of the kernel space&#xa0;in&#xa0;each convolutional layer. Additionally, ODConv outperforms other attention modules in adjusting output features or convolutional weights.</p>
<p>Considering that the Backbone network and Neck network of the YOLOv8n model mainly rely on the C2f module for connection, the main function of the C2f module is to input the multi-scale feature maps extracted by the Backbone network into the Neck network. To this end, the ODConv module is introduced into the C2f module that connects to the Neck network, which will fully exploit the target feature information across different dimensions. This allows richer image features to be transferred from the Backbone network to the Neck network, thereby enhancing the model&#x2019;s ability to recognize small and partially occluded targets. Subsequent experiments will show that not only is the model&#x2019;s performance significantly improved, but also the computational complexity is reduced during the training process.</p>
</sec>
<sec id="s2_5_2">
<label>2.5.2</label>
<title>ECA attention mechanism</title>
<p>The seedling soil used for hybrid rice seedling nursing is usually a mixture of peat, vermiculite, coconut coir, clay, specialized fertilizer, and yellow soil. Owing to the different colors and shapes of the components in the seedling soil, a complex background is formed in the seed images of hybrid rice, which interferes with the identification of the number of seeds in a unit seed grid image. In addition, the unit seed grid image is generated by cutting the seed image, and some seeds are cut into incomplete seeds, which also poses some difficulties for accurate seed counting. To improve the model&#x2019;s focus on seeds, this paper introduces the ECA attention mechanism (<xref ref-type="bibr" rid="B22">Wang et&#xa0;al., 2020</xref>) into the Neck network of the model (the introduced location is shown in <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>). This mechanism enables the network to prioritize the semantic information in the feature map, thereby enhancing the network&#x2019;s ability to extract deep semantic feature information of detection targets. Thus, the model can better obtain the feature information of seed images while suppressing background information such as seedling soil.</p>
<p>In neural networks, attention mechanisms are typically classified into spatial attention mechanisms and channel attention mechanisms. The channel attention mechanism assigns different weights to the channels of feature maps, allowing the model to have different degrees of attention to different channels of the feature map, thereby capturing local important feature information, suppressing interference feature information, and improving recognition accuracy. ECA is a type of channel attention mechanism that has been improved from the traditional Squeeze and Excitation Networks (SE) attention mechanism. The main improvements include (1) removing the fully connected layers in the SE attention mechanism module; (2) using a one-dimensional convolution to learn features after global average pooling (GAP); and (3) proposing a non-dimensional reduction local cross-channel interaction strategy, effectively avoiding the impact of dimensionality reduction on the learning effect of channel attention. Its module structure is shown in <xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8</bold>
</xref>.</p>
<fig id="f8" position="float">
<label>Figure&#xa0;8</label>
<caption>
<p>ECA module.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g008.tif"/>
</fig>
<p>In <xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8</bold>
</xref>, <italic>x</italic> represents the input features with dimensions <italic>H</italic>&#xd7;<italic>W</italic> &#xd7;<italic>C</italic>. After global average pooling (GAP), <italic>x</italic> is compressed into a feature vector of size 1&#xd7;1&#xd7;<italic>C</italic>. A one-dimensional convolution operation is performed on the feature vector, where the kernel size <italic>k</italic> is adaptively determined. The weights of each channel are computed through a Sigmoid activation function, denoted as <italic>&#x3c3;</italic> in <xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8</bold>
</xref>. Multiply the weighted feature vector with the input features <italic>x</italic> to get the output features <inline-formula>
<mml:math display="inline" id="im3">
<mml:mrow>
<mml:mover accent="true">
<mml:mi>x</mml:mi>
<mml:mo stretchy="true">&#x2dc;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula>, and in <xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8</bold>
</xref>, <inline-formula>
<mml:math display="inline" id="im4">
<mml:mo>&#x2297;</mml:mo>
</mml:math>
</inline-formula> denotes the multiplication.</p>
</sec>
<sec id="s2_5_3">
<label>2.5.3</label>
<title>WIoU loss function</title>
<p>YOLOv8n&#x2019;s loss function consists of two parts: classification loss and regression loss. The classification loss is computed using BCE Loss, while the regression loss is computed using CIoU Loss. The formula for calculating CIoU Loss is as follows:</p>
<disp-formula id="eq5">
<label>(5)</label>
<mml:math display="block" id="M5">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>CloU</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>+</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>+</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>y</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:msubsup>
<mml:mi>W</mml:mi>
<mml:mi>g</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mo>+</mml:mo>
<mml:msubsup>
<mml:mi>H</mml:mi>
<mml:mi>g</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mfrac>
<mml:mo>+</mml:mo>
<mml:mi>&#x3b1;</mml:mi>
<mml:mi>v</mml:mi>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq6">
<label>(6)</label>
<mml:math display="block" id="M6">
<mml:mrow>
<mml:mi>&#x3b1;</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mi>v</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>+</mml:mo>
<mml:mi>v</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq7">
<label>(7)</label>
<mml:math display="block" id="M7">
<mml:mrow>
<mml:mi>v</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mn>4</mml:mn>
<mml:mrow>
<mml:msup>
<mml:mi>&#x3c0;</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
<mml:msup>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mtext>arctan</mml:mtext>
<mml:mfrac>
<mml:mi>w</mml:mi>
<mml:mi>h</mml:mi>
</mml:mfrac>
<mml:mo>&#x2212;</mml:mo>
<mml:mtext>arctan</mml:mtext>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi>h</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq8">
<label>(8)</label>
<mml:math display="block" id="M8">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi>W</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>H</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:mi>w</mml:mi>
<mml:mi>h</mml:mi>
<mml:mo>+</mml:mo>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
<mml:msub>
<mml:mi>h</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>W</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>H</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<p>where <italic>w</italic>, <italic>h</italic>, and (<italic>x</italic>, <italic>y</italic>) represent the width, height, and center coordinates of the predicted bounding box, respectively; <italic>w<sub>gt</sub>
</italic>, <italic>h<sub>gt</sub>
</italic>, and (<italic>x<sub>gt</sub>
</italic>, <italic>y<sub>gt</sub>
</italic>) represent the width, height, and center coordinates of the ground truth bounding box, respectively; <italic>W<sub>g</sub>
</italic> and <italic>H<sub>g</sub>
</italic> represent the width and height of the smallest enclosing rectangle that contains both the ground truth and predicted bounding boxes; <italic>W<sub>i</sub>
</italic> and <italic>H<sub>i</sub>
</italic> denote the width and height of the intersection area between the ground truth and predicted bounding boxes. <italic>&#x3b1;</italic> is a weighting function used to balance the parameters, and <italic>v</italic> is a parameter used to measure the consistency of the length&#x2013;width ratio, as shown in <xref ref-type="fig" rid="f9">
<bold>Figure&#xa0;9</bold>
</xref>.</p>
<fig id="f9" position="float">
<label>Figure&#xa0;9</label>
<caption>
<p>Illustration of ground truth boxes, predicted boxes, and their intersection over union.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g009.tif"/>
</fig>
<p>Although the CIoU loss function considers overlapping area, distance, and aspect ratio, CIoU employs a monotonic focusing mechanism. This means that the bounding box regression loss uses fixed weights to balance the losses of position and scale. Because of the influence of geometric factors such as distance and aspect ratio, this can increase the penalty on low-quality samples and reduce the generalization performance of the YOLOv8n model. Specifically, when the aspect ratio of the predicted box and the ground truth box are linearly related, the penalty term of CIoU degrades to zero. In this case, whether the anchor boxes are of high quality or low quality, it can be detrimental to the regression loss (<xref ref-type="bibr" rid="B19">Tong et&#xa0;al., 2023</xref>).</p>
<p>Owing to the artificial cutting of some seeds in the unit seed grid image, there are some low-quality anchor boxes. Therefore, this paper chooses the dynamic non-monotonic focusing mechanism of the WIoU loss function to replace the original CIoU loss function. Essentially, the WIoU loss function adaptively adjusts the weights of regression loss based on the importance of the targets by introducing an attention mechanism. This helps to better balance the losses of position and scale, thereby balancing low-quality and high-quality samples to improve the accuracy of object detection. There are currently three versions of WIoU: WIoUv1, WIoUv2, and WIoUv3, among which WIoUv2 and WIoUv3 are enhanced versions of WIoUv1 (<xref ref-type="bibr" rid="B19">Tong et&#xa0;al., 2023</xref>). WIoUv3 is used in this paper, which is defined as follows:</p>
<disp-formula id="eq9">
<label>(9)</label>
<mml:math display="block" id="M9">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>WIoUv</mml:mtext>
<mml:mn>3</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mi>r</mml:mi>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>WIoUv</mml:mtext>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mi>&#x3b2;</mml:mi>
<mml:mrow>
<mml:mi>&#x3b4;</mml:mi>
<mml:msup>
<mml:mi>&#x3b1;</mml:mi>
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>&#x3b4;</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq10">
<label>(10)</label>
<mml:math display="block" id="M10">
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>L</mml:mi>
<mml:mo stretchy="true">&#x2dc;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#x2208;</mml:mo>
<mml:mo stretchy="false">[</mml:mo>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mo>+</mml:mo>
<mml:mi>&#x221e;</mml:mi>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq11">
<label>(11)</label>
<mml:math display="block" id="M11">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>WIoUv</mml:mtext>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:msub>
<mml:mi>R</mml:mi>
<mml:mrow>
<mml:mtext>WIoU</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq12">
<label>(12)</label>
<mml:math display="block" id="M12">
<mml:mrow>
<mml:msub>
<mml:mi>R</mml:mi>
<mml:mrow>
<mml:mi>W</mml:mi>
<mml:mi>I</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>U</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mtext>exp</mml:mtext>
<mml:mo>(</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>+</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>y</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:msubsup>
<mml:mi>W</mml:mi>
<mml:mi>g</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mo>+</mml:mo>
<mml:msubsup>
<mml:mi>H</mml:mi>
<mml:mi>g</mml:mi>
<mml:mn>2</mml:mn>
</mml:msubsup>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>*</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfrac>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:math>
</disp-formula>
<p>where <italic>r</italic> represents the non-monotonic focusing coefficient, <italic>&#x3b2;</italic> represents the outlier degree, and <italic>&#x3b1;</italic> and <italic>&#x3b4;</italic> represent hyperparameters, which are set to 1.9 and 3, respectively. <inline-formula>
<mml:math display="inline" id="im5">
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="true">&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula> denotes the moving average, and <inline-formula>
<mml:math display="inline" id="im6">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>L</mml:mi>
<mml:mo stretchy="true">&#x2dc;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> indicates an operation that separates from the computation graph, making it a constant without gradients. When <inline-formula>
<mml:math display="inline" id="im7">
<mml:mrow>
<mml:msub>
<mml:mi>R</mml:mi>
<mml:mrow>
<mml:mtext>WIoU</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:mo stretchy="false">[</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>,</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>, it increases the <italic>L</italic>
<sub>IoU</sub> of medium-quality anchor boxes, and when <inline-formula>
<mml:math display="inline" id="im8">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mtext>IoU</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:mo stretchy="false">[</mml:mo>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo stretchy="false">]</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>, it reduces the <italic>R</italic>
<sub>WIoU</sub> of high-quality anchor boxes.</p>
<p>WIoUv3 uses the outlier degree <italic>&#x3b2;</italic> instead of the overlap area IoU to evaluate the quality of the anchor box. Since the moving average is dynamic, the standard for classifying anchor box quality is also dynamic. The non-monotonic focusing coefficient <italic>r</italic> is constructed from the outlier degree and hyperparameters. It can reduce the gradient gain of the loss function on high-quality samples and decrease the harmful gradients generated by low-quality anchor boxes, allowing WIoUv3 to focus on medium-quality anchor boxes.</p>
</sec>
</sec>
<sec id="s2_6">
<label>2.6</label>
<title>Evaluation methods for sowing uniformity</title>
<sec id="s2_6_1">
<label>2.6.1</label>
<title>Determination of which seed grid the seed belongs to</title>
<p>After dividing the seed images into unit seed grid images, the seeds located on the boundaries of the seed grids may encounter the problem of misjudging the ownership of the seed grid.</p>
<p>Overall, there are four situations regarding the positional relationship between seeds and seed grids (as shown in <xref ref-type="fig" rid="f10">
<bold>Figure&#xa0;10</bold>
</xref>): (a) Seeds do not intersect with the boundaries of seed grids, and seeds completely fall within a certain seed grid; (b) the seed intersects with one seed grid boundary and falls into two adjacent seed grids; (c) the seeds intersect with the boundaries of two seed grids, and the seeds fall into three adjacent seed grids; (d) the seeds intersect with the boundaries of the four seed grids, and the seeds fall within the adjacent four seed grids.</p>
<fig id="f10" position="float">
<label>Figure&#xa0;10</label>
<caption>
<p>The positional relationship between seeds and seed grids. <bold>(A)</bold> the first state, <bold>(B)</bold> the second state, <bold>(C)</bold> the third state, and <bold>(D)</bold> the fourth state.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g010.tif"/>
</fig>
<p>At present, there is no good way to determine which seed grid a seed belongs to when there is a cross between the seed and the seed grid. In this paper, an area-based approach is employed to determine the ownership of seeds within seed grids. Specifically, the seed grid with the largest proportion of the area of the seed in the adjacent four seed grids is considered as the seed&#x2019;s belonging grid. The specific method is as follows: According to the OEW-YOLOv8n model provided in the following text, the seed recognition box covers the area size within four adjacent seed grids for discrimination, and the seed grid with the largest coverage area is the seed grid where the seed belongs.</p>
</sec>
<sec id="s2_6_2">
<label>2.6.2</label>
<title>Evaluation methods for uniformity</title>
<p>The key to evaluating the uniformity of seed images using uniform qualification rate is to calculate the number of qualified seed grid images for each seed image. According to the national standard of the People&#x2019;s Republic of China &#x201c;Rice Transplanter-Test Method&#x201d; (GB/T 6243-2017), the uniformity qualification rate is calculated as follows: the uniformity qualification rate = (number of qualified small seedling blocks/total number of seedling blocks tested) &#xd7; 100%. Because of the agronomic requirement of planting one to three seedlings per hill in hybrid rice cultivation, this article uses the qualification rate of one to three seeds per grid as the evaluation indicator of sowing uniformity. The qualification rate of one to three seeds per grid is calculated as: The qualification rate of one to three seeds per grid = (number of grids with one to three seeds per grid &#xf7; total number of seed grids per tray) &#xd7; 100%.</p>
</sec>
</sec>
</sec>
<sec id="s3" sec-type="results">
<label>3</label>
<title>Results</title>
<sec id="s3_1">
<label>3.1</label>
<title>Training environment and methods</title>
<p>Using the above method, we trained unit seed grid images. The dataset comprises a total of 23,814 images. Models such as OEW-YOLOv8n were used, with the input image resolution adjusted to 640 &#xd7; 640 pixels. In order to adapt to the convergence effect of the model and the complexity of the dataset, The optimizer employed for training was Stochastic Gradient Descent (SGD). The number of workers was set to 2, and the batch size was set to 32. The initial learning rate was 0.01, and the training lasted for 200 epochs. The model training and testing environment included a 64-bit Windows 10 Professional operating system, a 13<sup>th</sup>-generation Intel(R) Core (TM) i5-13600KF CPU running at 3.5 GHz, 32 GB of RAM, an NVIDIA GeForce RTX 2070 SUPER GPU with 8 GB of VRAM, CUDA version 11.7, the deep learning framework PyTorch1.12, and Python version 3.8.</p>
</sec>
<sec id="s3_2">
<label>3.2</label>
<title>Evaluating indicator</title>
<p>In this research, Params, floating point operations (FLOPs), and model size are used to reflect the complexity of a model, where FLOPs are used to measure computational amount. Precision (<italic>P</italic>), recall (<italic>R</italic>), and mAP [as described in <xref ref-type="disp-formula" rid="eq13">Equations 13</xref>&#x2013;<xref ref-type="disp-formula" rid="eq16">16</xref>] are used to evaluate the detection performance of a model. Root mean square error (RMSE) is used to quantify the deviation between the model&#x2019;s estimated value and true value, with a lower score indicating less error and better model performance.</p>
<disp-formula id="eq13">
<label>(13)</label>
<mml:math display="block" id="M13">
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
</mml:msub>
<mml:mo>+</mml:mo>
<mml:msub>
<mml:mi>F</mml:mi>
<mml:mi>p</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>100</mml:mn>
<mml:mo>%</mml:mo>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq14">
<label>(14)</label>
<mml:math display="block" id="M14">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
</mml:msub>
<mml:mo>+</mml:mo>
<mml:msub>
<mml:mi>F</mml:mi>
<mml:mi>N</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>100</mml:mn>
<mml:mo>%</mml:mo>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq15">
<label>(15)</label>
<mml:math display="block" id="M15">
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>=</mml:mo>
<mml:mstyle displaystyle="true">
<mml:mrow>
<mml:msubsup>
<mml:mo>&#x222b;</mml:mo>
<mml:mn>0</mml:mn>
<mml:mn>1</mml:mn>
</mml:msubsup>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mtext>d</mml:mtext>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:mrow>
</mml:mstyle>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>100</mml:mn>
<mml:mo>%</mml:mo>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq16">
<label>(16)</label>
<mml:math display="block" id="M16">
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:math>
</disp-formula>
<p>In <xref ref-type="disp-formula" rid="eq13">Equations 13</xref>&#x2013;<xref ref-type="disp-formula" rid="eq16">16</xref>, <italic>T<sub>p</sub>
</italic> and <italic>F<sub>p</sub>
</italic> represent true-positive and false-positive instances, respectively. <italic>F<sub>n</sub>
</italic> represents false-negative instances. <italic>AP</italic>(<italic>i</italic>) denotes the average precision for the <italic>i</italic>th class, and <italic>N</italic> represents the total number of classes. In this paper, <italic>N</italic> equals 1, as there is only one class, which is the seed class.</p>
<disp-formula id="eq17">
<label>(17)</label>
<mml:math display="block" id="M17">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>E</mml:mi>
<mml:mo>=</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mi>M</mml:mi>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>M</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:msubsup>
<mml:mstyle displaystyle="true">
<mml:mi>y</mml:mi>
</mml:mstyle>
<mml:mi>i</mml:mi>
<mml:mrow>
<mml:mi>p</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x2212;</mml:mo>
<mml:msubsup>
<mml:mstyle displaystyle="true">
<mml:mi>y</mml:mi>
</mml:mstyle>
<mml:mi>i</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:math>
</disp-formula>
<p>where <inline-formula>
<mml:math display="inline" id="im9">
<mml:mrow>
<mml:msubsup>
<mml:mstyle displaystyle="true">
<mml:mi>y</mml:mi>
</mml:mstyle>
<mml:mi>i</mml:mi>
<mml:mrow>
<mml:mi>p</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> represents the estimated value, <inline-formula>
<mml:math display="inline" id="im10">
<mml:mrow>
<mml:msubsup>
<mml:mstyle displaystyle="true">
<mml:mi>y</mml:mi>
</mml:mstyle>
<mml:mi>i</mml:mi>
<mml:mrow>
<mml:mtext>g</mml:mtext>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> represents the true value, and <italic>M</italic> represents the number of samples.</p>
</sec>
<sec id="s3_3">
<label>3.3</label>
<title>The detection results of seed grid and seed count based on the OEW-YOLOv8n model</title>
<sec id="s3_3_1">
<label>3.3.1</label>
<title>Ablation experiments</title>
<p>The results of the ablation experiment are shown in <xref ref-type="table" rid="T3">
<bold>Table&#xa0;3</bold>
</xref>.</p>
<table-wrap id="T3" position="float">
<label>Table&#xa0;3</label>
<caption>
<p>Ablation experiment results.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="center">Number</th>
<th valign="middle" align="center">ODConv</th>
<th valign="middle" align="center">ECA</th>
<th valign="middle" align="center">WIoU</th>
<th valign="middle" align="center">Params (M)</th>
<th valign="middle" align="center">FLOPs (G)</th>
<th valign="middle" align="center">Model size (MB)</th>
<th valign="middle" align="center">P (%)</th>
<th valign="middle" align="center">R (%)</th>
<th valign="middle" align="center">mAP@0.5 (%)</th>
<th valign="middle" align="center">RMSE</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="center">0</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">3.01</td>
<td valign="middle" align="center">8.2</td>
<td valign="middle" align="center">6.13</td>
<td valign="middle" align="center">92.3</td>
<td valign="middle" align="center">90.5</td>
<td valign="middle" align="center">96.1</td>
<td valign="middle" align="center">0.445</td>
</tr>
<tr>
<td valign="middle" align="center">1</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">3.35</td>
<td valign="middle" align="center">8.0</td>
<td valign="middle" align="center">6.83</td>
<td valign="middle" align="center">95.1</td>
<td valign="middle" align="center">93.6</td>
<td valign="middle" align="center">97.4</td>
<td valign="middle" align="center">0.392</td>
</tr>
<tr>
<td valign="middle" align="center">2</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">3.01</td>
<td valign="middle" align="center">8.2</td>
<td valign="middle" align="center">6.14</td>
<td valign="middle" align="center">94.2</td>
<td valign="middle" align="center">93.1</td>
<td valign="middle" align="center">96.7</td>
<td valign="middle" align="center">0.426</td>
</tr>
<tr>
<td valign="middle" align="center">3</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">3.01</td>
<td valign="middle" align="center">8.2</td>
<td valign="middle" align="center">6.13</td>
<td valign="middle" align="center">93.8</td>
<td valign="middle" align="center">92.7</td>
<td valign="middle" align="center">96.5</td>
<td valign="middle" align="center">0.435</td>
</tr>
<tr>
<td valign="middle" align="center">4</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">3.35</td>
<td valign="middle" align="center">8.0</td>
<td valign="middle" align="center">6.84</td>
<td valign="middle" align="center">96.2</td>
<td valign="middle" align="center">94.9</td>
<td valign="middle" align="center">98.1</td>
<td valign="middle" align="center">0.299</td>
</tr>
<tr>
<td valign="middle" align="center">5</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">3.35</td>
<td valign="middle" align="center">8.0</td>
<td valign="middle" align="center">6.83</td>
<td valign="middle" align="center">95.2</td>
<td valign="middle" align="center">94.5</td>
<td valign="middle" align="center">97.7</td>
<td valign="middle" align="center">0.365</td>
</tr>
<tr>
<td valign="middle" align="center">6</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">3.01</td>
<td valign="middle" align="center">8.2</td>
<td valign="middle" align="center">6.14</td>
<td valign="middle" align="center">94.3</td>
<td valign="middle" align="center">93.2</td>
<td valign="middle" align="center">97.2</td>
<td valign="middle" align="center">0.417</td>
</tr>
<tr>
<td valign="middle" align="center">7</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">3.35</td>
<td valign="middle" align="center">8.0</td>
<td valign="middle" align="center">6.84</td>
<td valign="middle" align="center">96.8</td>
<td valign="middle" align="center">95.3</td>
<td valign="middle" align="center">98.6</td>
<td valign="middle" align="center">0.269</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>&#x201c;&#x221a;&#x201d; indicates that the corresponding policy is used, and &#x201c;&#xd7;&#x201d; indicates that the corresponding policy is not used.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>Based on <xref ref-type="table" rid="T3">
<bold>Table&#xa0;3</bold>
</xref>, it is evident that after replacing the Conv module in the YOLOv8n network with the ODConv module in Experiment 1, the Params and model size slightly increased, but the FLOPs decreased. The precision (<italic>P</italic>), recall (<italic>R</italic>), and mAP of the model increased by 2.8%, 3.1%, and 1.3%, respectively. The RMSE decreased by 12%. Experiments 2 and 3 showed that adding ECA attention mechanism or replacing CIoU loss function with the WIoUv3 loss function did not significantly change the Parms, FLOPs, and model size. The precision of the model increased by 1.9% and 1.5%, the recall increased by 2% and 2.2%, and the mAP increased by 0.6% and 0.4%, respectively. The RMSE decreased by 4% and 2%. These experiments demonstrate that both the ECA attention mechanism and the WIoUv3 loss function can enhance the model&#x2019;s detection accuracy. Experiment 4 is an experiment adding the ECA attention mechanism on Experiment 1. The Params and FLOPs remained nearly unchanged, while the model size slightly increased. The precision, recall, and mAP further improved, with the mAP increasing from 97.4% to 98.1%, an increment of 0.7%. The RMSE decreased by 23.7%. Experiment 5: Building on Experiment 1 by replacing the CIoU loss function with the WIoUv3 loss function, there was no significant change in the Params, FLOPs, and model size. However, precision, recall, and mAP improved by 0.1%, 0.9%, and 0.3%, respectively. The RMSE decreased by 7%. Experiment 6: Building on Experiment 2 by replacing the CIoU loss function with the WIoUv3 loss function, the Params, FLOPs, and model size remained largely unchanged, with improvements observed in precision, recall, and mAP. The RMSE decreased by 2%. Experiment 7: Building on Experiment 4 by replacing the CIoU loss function with the WIoUv3 loss function, the Params, FLOPs, and model size remained largely unchanged. Precision, recall, and mAP improved by 0.6%, 0.4%, and 0.5%, respectively. The RMSE decreased by 10%.</p>
<p>In summary, using the ODConv module to replace the original Conv module slightly increases the Params and model size, but reduces the FLOPs by 0.2 G and improves the mAP by 1.3%. Adding the ECA attention mechanism does not significantly change the model&#x2019;s complexity, yet it raises the model&#x2019;s mAP by approximately 0.6%. Replacing the CIoU loss function with the WIoUv3 loss function also does not alter the model&#x2019;s complexity, but the precision improves by 0.4%. Implementing all three strategies simultaneously results in an increase of 0.34 M in Params, 0.71 MB in model size, a reduction of 0.2 G in FLOPs, and improvements in precision, recall, and mAP by 4.5%, 4.8%, and 2.5%, respectively. This indicates that all three improvement strategies enhance the model&#x2019;s detection accuracy and have a cumulative effect, with the ODConv having the most significant impact. The slight increase in Params and model size is primarily due to the introduction of the ODConv module.</p>
<p>To examine the impact of different strategies on the convergence speed of the model, <xref ref-type="fig" rid="f11">
<bold>Figure&#xa0;11</bold>
</xref> shows the curve of mAP of different models in the ablation experiment as the number of model iterations changes.</p>
<fig id="f11" position="float">
<label>Figure&#xa0;11</label>
<caption>
<p>The curve of mAP.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g011.tif"/>
</fig>
<p>From <xref ref-type="fig" rid="f11">
<bold>Figure&#xa0;11</bold>
</xref>, it can be found that after approximately 25 iterations, the fluctuations in the mAP curve begin to diminish, and after approximately 150 iterations, the mAP curve tends to stabilize. This indicates that all models exhibit good convergence speed. Among them, Experiment 7 performs the best, while Experiment 0 performs the worst. This demonstrates that all three improvement strategies enhance convergence speed, and employing all three strategies simultaneously can achieve relatively faster convergence speeds and higher convergence values.</p>
</sec>
<sec id="s3_3_2">
<label>3.3.2</label>
<title>Comparisons of OEW-YOLOv8n performance under different loss functions</title>
<p>In this study, the bounding box regression loss function has a significant impact on the accurate identification of the number of seeds in a unit seed grid. In existing deep learning models, common regression loss functions include CIoU, EIoU, SIoU, and WIoU. The loss function of the YOLOv8n model is CIoU. In order to investigate the influence of different regression loss functions on the performance of the OEW-YOLOv8n model, the convergence of OEW-YOLOv8n under different loss functions was compared and analyzed. The experimental results are illustrated in <xref ref-type="fig" rid="f12">
<bold>Figure&#xa0;12</bold>
</xref>.</p>
<fig id="f12" position="float">
<label>Figure&#xa0;12</label>
<caption>
<p>Bounding box loss.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g012.tif"/>
</fig>
<p>As shown in <xref ref-type="fig" rid="f12">
<bold>Figure&#xa0;12</bold>
</xref>, when the bounding box regression loss function adopts EIoU, the convergence speed of the model is the slowest, and the final loss value after convergence is the highest. When using the CIoU loss function, its convergence speed and final loss value are almost equivalent to EIoU (shown in the graph as the red and gray lines overlapping). The convergence speed of using SIoU is slightly higher than that of EIoU and CIoU, and the final loss value after convergence is slightly lower than EIoU and CIoU. When WIoU is adopted, the convergence speed of the model is the fastest, converging approximately after 25 iterations. Although there are some fluctuations in the final loss value after convergence, which may be related to the use of non-monotonic focusing mechanism, the final loss value after convergence is significantly lower than those of the other three loss functions. The ablation experiment in Section 3.3.1 also showed that the use of WIoU resulted in some improvement in mAP. Therefore, this paper adopts WIoU to replace the CIoU loss function.</p>
</sec>
<sec id="s3_3_3">
<label>3.3.3</label>
<title>Model comparison</title>
<p>To evaluate the performance of the OEW-YOLOv8n model, comparative experiments were conducted with advanced object detection networks such as Faster-RCNN, SSD, YOLOv4, YOLOv5s, YOLOv7-tiny, and YOLOv10s. The dataset used in these experiments is the seed grid images used in this study. The experimental results are shown in <xref ref-type="table" rid="T4">
<bold>Table&#xa0;4</bold>
</xref>.</p>
<table-wrap id="T4" position="float">
<label>Table&#xa0;4</label>
<caption>
<p>Comparison of seed number detection results of different models.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="left">Models</th>
<th valign="middle" align="left">mAP@0.5 (%)</th>
<th valign="middle" align="left">Params (M)</th>
<th valign="middle" align="left">FLOPs (G)</th>
<th valign="middle" align="left">Model size (MB)</th>
<th valign="middle" align="left">RMSE</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="center">Faster-RCNN</td>
<td valign="middle" align="center">93.4</td>
<td valign="middle" align="center">136.68</td>
<td valign="middle" align="center">199.3</td>
<td valign="middle" align="center">110.76</td>
<td valign="middle" align="center">0.600</td>
</tr>
<tr>
<td valign="middle" align="center">SSD</td>
<td valign="middle" align="center">90.8</td>
<td valign="middle" align="center">23.61</td>
<td valign="middle" align="center">273.2</td>
<td valign="middle" align="center">92.78</td>
<td valign="middle" align="center">0.733</td>
</tr>
<tr>
<td valign="middle" align="center">YOLOv4</td>
<td valign="middle" align="center">93.7</td>
<td valign="middle" align="center">63.94</td>
<td valign="middle" align="center">59.9</td>
<td valign="middle" align="center">250.25</td>
<td valign="middle" align="center">0.590</td>
</tr>
<tr>
<td valign="middle" align="center">YOLOv5s</td>
<td valign="middle" align="center">95.8</td>
<td valign="middle" align="center">6.83</td>
<td valign="middle" align="center">15.9</td>
<td valign="middle" align="center">14.73</td>
<td valign="middle" align="center">0.514</td>
</tr>
<tr>
<td valign="middle" align="center">YOLOv7-tiny</td>
<td valign="middle" align="center">95.7</td>
<td valign="middle" align="center">6.02</td>
<td valign="middle" align="center">13.2</td>
<td valign="middle" align="center">12.32</td>
<td valign="middle" align="center">0.544</td>
</tr>
<tr>
<td valign="middle" align="center">YOLOv10s</td>
<td valign="middle" align="center">95.3</td>
<td valign="middle" align="center">7.19</td>
<td valign="middle" align="center">21.6</td>
<td valign="middle" align="center">14.69</td>
<td valign="middle" align="center">0.567</td>
</tr>
<tr>
<td valign="middle" align="center">OEW-YOLOv8n</td>
<td valign="middle" align="center">98.6</td>
<td valign="middle" align="center">3.35</td>
<td valign="middle" align="center">8.0</td>
<td valign="middle" align="center">6.84</td>
<td valign="middle" align="center">0.269</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>According to <xref ref-type="table" rid="T4">
<bold>Table&#xa0;4</bold>
</xref>, the Params, FLOPs, model size, and RMSE of the OEW-YOLOv8n model are all smaller than those of the other five models. Compared to Faster-RCNN, SSD, YOLOv4, YOLOv5s, YOLOv7-tiny, and YOLOv10s, the Params of the OEW-YOLOv8n model is reduced by 133.33 M, 20.26 M, 60.59 M, 3.48 M, 2.67 M, and 3.84 M, respectively. The FLOPs are reduced by 191.3 G, 265.2 G, 51.9 G, 7.9 G, 5.2 G, and 13.6 G, respectively, and the model size is reduced by 103.92 MB, 85.94 MB, 243.41 MB, 6.89 MB, 5.19 MB, and 7.85 MB, respectively. These indicate that the complexity of the model has significantly decreased. Conversely, the mAP of the OEW-YOLOv8n model is higher than that of the other five models, with improvements of 5.2%, 7.8%, 4.9%, 2.8%, 2.9%, and 3.3% over Faster-RCNN, SSD, YOLOv4, YOLOv5s, YOLOv7-tiny, and YOLOv10s, respectively. The RMSE is reduced by 0.331, 0.464, 0.321, 0.245, 0.275, and 0.298, respectively. Overall, the two-stage network model Faster-RCNN has high model complexity and low recognition accuracy. The one-stage network models such as SSD, YOLOv4, YOLOv5s, YOLOv7-tiny, and YOLOv10s have a relatively lower model complexity and a higher recognition accuracy.</p>
<p>To analyze the reasons for the differences in detection performance among different object detection models, <xref ref-type="fig" rid="f13">
<bold>Figure&#xa0;13</bold>
</xref> shows the detection results of a representative seed grid under different sowing densities and sowing methods. In the experiment, to increase the number of seeds and reduce the segmentation of seeds during seed grid division and test the generalization ability of the model, a seed image was divided into 144 seed grids according to a specification of 9 times by 16 times in the experiment.</p>
<fig id="f13" position="float">
<label>Figure&#xa0;13</label>
<caption>
<p>Comparison of detection performance of different object detection models. S1&#x2013;S7 represent the SSD, Faster-RCNN, YOLOv4, YOLOv5s, YOLOv7-tiny, YOLOv10s, and OEW-YOLOv8 algorithms, respectively, while B1&#x2013;B9 represent broadcast sowing with 50 g/tray, no-ditching drill sowing with 50 g/tray, ditching drill sowing with 50 g/tray, broadcast sowing with 60 g/tray, no-ditching drill sowing with 60 g/tray, ditching drill sowing with 60 g/tray, broadcast sowing with 70 g/tray, no-ditching drill sowing with 70 g/tray, and ditching drill sowing with 70 g/tray, respectively.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1473153-g013.tif"/>
</fig>
<p>From <xref ref-type="fig" rid="f13">
<bold>Figure&#xa0;13</bold>
</xref>, it can be seen that the SSD model performs the worst, with a large number of missed detections. For example, in cell B3, there are 13 seeds in the image (with 11 actual effective seeds), but the SSD model only identifies 8 seeds, missing 5 seeds, resulting in a 38.46% miss rate. In the SSD model, the missed seeds are mainly defective seeds at the edges of the image or seeds that are obscured. While Faster-RCNN, YOLOv4, YOLOv5s, YOLOv7-tiny, and YOLOv10s can identify defective seeds at the image edges, they also exhibit missed detections and have poor recognition accuracy for seeds obscured by seedling soil. The OEW-YOLOv8n model performs very well overall, capable of identifying both defective seeds at the image edges and obscured seeds, with almost no missed detections.</p>
<p>In summary, compared to the existing advanced networks, the improved OEW-YOLOv8n model performs better in detecting the number of seeds per unit seed grid.</p>
</sec>
</sec>
<sec id="s3_4">
<label>3.4</label>
<title>Evaluation experiment of sowing uniformity in hybrid rice blanket-seedling nursing</title>
<p>To investigate the practical application effect of the algorithm on the evaluation of sowing uniformity, a comparative experiment was conducted between the evaluation of sowing uniformity algorithm and manual evaluation, with uniform qualification rate and empty grid rate as evaluation indicators. The evaluation algorithm adopts the method described in Section 2 of this article. The manual evaluation adopts the method of manually counting the number of seeds per unit grid, as follows:</p>
<p>According to the method described in <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref>, seedling trays with 50 g/tray, 60 g/tray, and 70 g/tray were divided into 738 grids/tray, 864 grids/tray, and 1,044 grids/tray, respectively. In the experiment, the grid frame was made of transparent thin wires, and the grid frame specifications for 50 g/tray, 60 g/tray, and 70 g/tray were 15.56 mm &#xd7; 14.15 mm, 15.56 mm &#xd7; 12.08 mm, and 15.56 mm &#xd7; 10.00 mm, respectively. These grid frames were used to divide the seed trays into seed grids. Then, the number of seeds in each seed grid was counted. The uniformity qualification rate was calculated according to the method in Section 2.6.2, and the empty grid rate was calculated as follows: empty grid rate = (number of empty grid/total number of grids) &#xd7; 100%. For each type of seedling trays, three trays were tested, which is equivalent to three repetitions. The average of the results of the three tests was taken as the corresponding test result. The results are shown in <xref ref-type="table" rid="T5">
<bold>Table&#xa0;5</bold>
</xref>.</p>
<table-wrap id="T5" position="float">
<label>Table&#xa0;5</label>
<caption>
<p>Sowing uniformity evaluation test results.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" rowspan="2" align="center">Seedling density/(g&#xb7;tray&#x2212;1)</th>
<th valign="middle" colspan="3" align="center">Seedling type</th>
<th valign="middle" rowspan="2" align="center">Testing empty grid rate (%)</th>
<th valign="middle" rowspan="2" align="center">Actual empty grid rate (%)</th>
<th valign="middle" rowspan="2" align="center">Testing <break/>uniformity <break/>qualification rate (%)</th>
<th valign="middle" rowspan="2" align="center">Actual <break/>uniformity <break/>qualification rate (%)</th>
<th valign="middle" rowspan="2" align="center">Test error (%)</th>
</tr>
<tr>
<th valign="middle" align="center">Broadcast sowing</th>
<th valign="middle" align="center">No-ditching <break/>strip sowing</th>
<th valign="middle" align="center">Ditching strip sowing</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" rowspan="3" align="center">50</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">4.67</td>
<td valign="middle" align="center">4.82</td>
<td valign="middle" align="center">68.75</td>
<td valign="middle" align="center">70.88</td>
<td valign="middle" align="center">2.13</td>
</tr>
<tr>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">4.30</td>
<td valign="middle" align="center">4.26</td>
<td valign="middle" align="center">72.63</td>
<td valign="middle" align="center">75.47</td>
<td valign="middle" align="center">2.84</td>
</tr>
<tr>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">3.75</td>
<td valign="middle" align="center">3.54</td>
<td valign="middle" align="center">82.77</td>
<td valign="middle" align="center">80.82</td>
<td valign="middle" align="center">&#x2212;1.95</td>
</tr>
<tr>
<td valign="middle" rowspan="3" align="center">60</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">3.82</td>
<td valign="middle" align="center">3.79</td>
<td valign="middle" align="center">70.83</td>
<td valign="middle" align="center">69.27</td>
<td valign="middle" align="center">&#x2212;1.56</td>
</tr>
<tr>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">3.19</td>
<td valign="middle" align="center">3.05</td>
<td valign="middle" align="center">73.05</td>
<td valign="middle" align="center">75.10</td>
<td valign="middle" align="center">2.05</td>
</tr>
<tr>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">2.77</td>
<td valign="middle" align="center">2.72</td>
<td valign="middle" align="center">77.91</td>
<td valign="middle" align="center">80.83</td>
<td valign="middle" align="center">2.92</td>
</tr>
<tr>
<td valign="middle" rowspan="3" align="center">70</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">2.75</td>
<td valign="middle" align="center">2.62</td>
<td valign="middle" align="center">65.97</td>
<td valign="middle" align="center">67.22</td>
<td valign="middle" align="center">1.25</td>
</tr>
<tr>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">1.94</td>
<td valign="middle" align="center">2.04</td>
<td valign="middle" align="center">69.86</td>
<td valign="middle" align="center">72.43</td>
<td valign="middle" align="center">2.57</td>
</tr>
<tr>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#xd7;</td>
<td valign="middle" align="center">&#x221a;</td>
<td valign="middle" align="center">1.11</td>
<td valign="middle" align="center">1.18</td>
<td valign="middle" align="center">79.72</td>
<td valign="middle" align="center">77.29</td>
<td valign="middle" align="center">&#x2212;2.43</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>In <xref ref-type="table" rid="T5">
<bold>Table&#xa0;5</bold>
</xref>, the empty grid rate and uniformity qualification rate resulting from algorithm evaluation are respectively named as testing empty grid rate and testing uniformity qualification rate, and the empty grid rate and uniformity qualification rate resulting from manual evaluation are named actual empty grid rate and actual uniformity qualification rate, respectively. Test error = actual uniformity qualification rate &#x2212; testing uniformity qualification rate.</p>
<p>According to <xref ref-type="table" rid="T5">
<bold>Table&#xa0;5</bold>
</xref>, the experiment indicates errors of &#x2212;0.15% to 0.21% for empty grid rate and &#x2212;2.43% to 2.92% detection, indicating that the algorithm has good accuracy. The actual uniformity qualification rates of ditching strip sowing with 50 g/tray, 60 g/tray, and 70 g/tray are 80.82%, 80.83%, and 77.29%, respectively, with an average of 79.65%. The actual uniformity qualification rates of no-ditching drill sowing with 50 g/tray, 60 g/tray, and 70 g/tray are 75.47%, 75.10%, and 72.43%, respectively, with an average of 74.33%. The actual uniformity qualification rates of broadcast sowing with 50 g/tray, 60 g/tray, and 70 g/tray are 70.88%, 69.27%, and 67.22%, respectively, with an average of 69.12%. From this, it can be seen that the effect of sowing density on sowing uniformity is relatively small when the sowing density is not too high. Among different sowing methods, the order of uniformity qualification rates from high to low is as follows: ditching strip sowing &gt; no-ditching strip drill sowing &gt; broadcast sowing, which is consistent with the actual situation.</p>
</sec>
</sec>
<sec id="s4" sec-type="discussion">
<label>4</label>
<title>Discussion</title>
<p>(1) The evaluation method proposed in this paper is mainly suitable for evaluating the sowing uniformity of rice under low-density conditions (usually referring to sowing density less than 70 g/tray). On one hand, according to the agronomic requirements of rice production (five to seven seedlings per hill for conventional rice and one to three seedlings per hill for hybrid rice), high-density sowing is usually used for conventional rice seedling nursing, while low-density sowing is used for hybrid rice seedling nursing. When high-density sowing is used for seedling nursing, the impact of sowing uniformity on the quality of rice transplanting is relatively small; however, when using low-density sowing for seedling nursery, the sowing uniformity has a significant impact on the quality of rice transplanting. Therefore, the evaluation of rice sowing uniformity is mainly applied in low-density sowing for hybrid rice seedling nursing. On the other hand, when the sowing density is very high, such as more than 120 g/tray, there will be a lot of overlapping seeds in the seedling tray. It is very difficult to evaluate sowing uniformity whether manually or through machine vision, and its evaluation significance is also very limited. As a result, this paper does not involve the evaluation in high-density sowing condition, and the impact of seed overlapping on the algorithm has not been considered temporarily.</p>
<p>(2) Assuming that one grid in a blanket tray is equivalent to one pot in a pot tray, we can find that the detection method adopted by <xref ref-type="bibr" rid="B18">Tan et&#xa0;al. (2014)</xref> achieved a mAP 94.4%, while our research method achieved 98.6%, an improvement of 4.2 percentage points. <xref ref-type="bibr" rid="B6">Dong&#xa0;et&#xa0;al. (2020)</xref> proposed a detection method with an mAP of 90.98%, whereas our research method achieved 98.6%, an improvement of 7.62 percentage points.</p>
<p>(3) In the detection area of this article, a seedling tray includes 738&#x2013;1,044 grids. In terms of detection time, the average detection time is 7 ms per unit seed grid image. Compared to the detection method proposed by <xref ref-type="bibr" rid="B18">Tan et&#xa0;al. (2014)</xref>, the detection time per grid image is reduced from 10.28 to 7 ms, a decrease of 3.28 ms.</p>
<p>(4) In practical applications, the detection of sowing uniformity is mainly applied in two places: first, evaluating the performance testing of seeders carried out by institutions; the second is to adjust the performance of the seeder during production to ensure the stability of sowing quality. Objectively speaking, the method proposed in this article can adapt well to the first scenario, but it still cannot adapt well to the second scenario. There are still some difficulties in deploying this method to achieve real-time detection on edge devices such as smartphones. This is a limitation of the evaluation method proposed in this paper. In the future, we will research how to achieve real-time online detection without reducing detection accuracy. In addition, to meet the needs of large-scale farm production, we will conduct systematic research on the impact&#xa0;of factors such as lighting environment, variety, and image&#xa0;acquisition equipment on the detection effect in our subsequent work.</p>
</sec>
<sec id="s5" sec-type="conclusions">
<label>5</label>
<title>Conclusions</title>
<p>In response to the issue of evaluation of sowing uniformity in hybrid rice blanket-seedling nursing, this paper proposes an evaluation method that integrates image processing with the OEW-YOLOv8n model, and an evaluation experiment for the sowing uniformity in hybrid rice blanket-seedling nursing was conducted. The following conclusions were mainly obtained:</p>
<list list-type="simple">
<list-item>
<p>(1) On the test set, the model achieves a mAP of 98.6%. The model size is 6.84 MB, the FLOPs are 8.0 G, and the number of parameters is 3.35 M. Compared to the original model, with a slight increase in model size and parameters, the FLOPs, precision, recall, and mAP increased by 0.2 G, 4.5%, 4.8%, and 2.5%, respectively.</p>
</list-item>
<list-item>
<p>(2) Among the three improvement strategies suggested, the strategy of using ODConv module has the greatest impact on model improvement. This strategy can respectively increase the precision, recall, and mAP by 2.8%, 3.1%, and 1.3%. Similarly, adding the ECA attention mechanism or replacing the CIoU loss function with the WIoUv3 loss function can increase the precision, recall, and mAP by 1.9%, 2.6%, and 0.6%, or 1.5%, 2.2%, and 0.4%, respectively. In addition, replacing the CIoU loss function with the WIoUv3 loss function has little effect on the complexity of the model, and these three improvement strategies have a cumulative effect on improving the detection performance of the model.</p>
</list-item>
<list-item>
<p>(3) The analysis of application sample experiments shows that the algorithm proposed in this paper has satisfactory accuracy. The experiment indicates errors of &#x2212;0.15% to 0.21% for empty grid rate and &#x2212;2.43% to 2.92% detection. Furthermore, the study also found that the sowing density has little impact on sowing uniformity in low-density sowing. In terms of different sowing methods, the order of sowing uniformity qualification rates from high to low is&#xa0;ditching strip sowing &gt; no-ditching strip sowing &gt;&#xa0;broadcast sowing, which is consistent with the actual situations.</p>
</list-item>
</list>
</sec>
</body>
<back>
<sec id="s6" sec-type="data-availability">
<title>Data availability statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec id="s7" sec-type="author-contributions">
<title>Author contributions</title>
<p>ZL: Writing &#x2013; original draft, Writing &#x2013; review &amp; editing, Funding acquisition, Project administration, Visualization. YP: Writing &#x2013; original draft, Writing &#x2013; review &amp; editing, Investigation, Methodology, Validation. XM: Writing &#x2013; review &amp; editing, Conceptualization, Formal analysis, Funding acquisition. YL: Writing &#x2013; review &amp; editing, Resources. XW: Writing &#x2013; review &amp; editing, Software. HL: Writing &#x2013; review &amp; editing, Supervision.</p>
</sec>
<sec id="s8" sec-type="funding-information">
<title>Funding</title>
<p>The author(s) declare financial support was received for the&#xa0;research, authorship, and/or publication of this article. This work was supported by the National Natural Science Foundation of China (No. 52175226) and the Research and Development Projects&#xa0;in Key Areas of Guangdong (No. 2023B0202130001).</p>
</sec>
<ack>
<title>Acknowledgments</title>
<p>We are very grateful to South China Agricultural University for supporting this study.</p>
</ack>
<sec id="s9" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s10" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Alekseev</surname> <given-names>E. P.</given-names>
</name>
<name>
<surname>Vasiliev</surname> <given-names>S. A.</given-names>
</name>
<name>
<surname>Maksimov</surname> <given-names>I. I.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Investigation of seed uniformity under field and laboratory conditions</article-title>. <source>IOP Conf. Series: Mater. Sci. Eng.</source> <volume>450</volume>, <elocation-id>62011</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1088/1757-899X/450/6/062011</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Ding</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Gong</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Lian</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Research on prediction of seed number in the hole based on machine vision and GABP algorithm</article-title>. <source>Measurement Control Technol.</source> <volume>36</volume>, <fpage>18</fpage>&#x2013;<lpage>23</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.19708/j.ckjs.2017.09.004</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Qiang</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Sugarcane nodes identification algorithm based on sum of local pixel of minimum points of vertical projection function</article-title>. <source>Comput. Electron. Agric.</source> <volume>182</volume>, <fpage>105994</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compag.2021.105994</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Lei</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Tang</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Deng</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Xiao</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Challenges and strategies of hybrid rice development in China</article-title>. <source>Hybrid Rice</source> <volume>30</volume>, <fpage>1</fpage>&#x2013;<lpage>4</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.16267/j.cnki.1005-3956.201505001</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Dai</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Yuan</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>Dynamic convolution: Attention over convolution kernels</article-title>,&#x201d; in <conf-name>2020 IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR).</conf-name> (<publisher-loc>Washington</publisher-loc>: <publisher-name>IEEE Press</publisher-name>), <fpage>11027</fpage>&#x2013;<lpage>11036</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/CVPR42600.2020.01104</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dong</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Tan</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Guo</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Design of low seeding quantity detection and control device for hybrid rice utilizing embedded machine vision</article-title>. <source>J. Jilin University(Engineering Technol. Edition)</source> <volume>50</volume>, <fpage>2295</fpage>&#x2013;<lpage>2305</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.13229/j.cnki.jdxbgxb20190818</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jiang</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Johansen</surname> <given-names>K. L.</given-names>
</name>
<name>
<surname>Stanschewski</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Wellman</surname> <given-names>G. B.</given-names>
</name>
<name>
<surname>Mousa</surname> <given-names>M. A. A.</given-names>
</name>
<name>
<surname>Fiene</surname> <given-names>G.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Phenotyping a diversity panel of quinoa using UAV-retrieved leaf area index, SPAD-based chlorophyll and a random forest approach</article-title>. <source>Precis. Agric.</source> <volume>23</volume>, <fpage>961</fpage>&#x2013;<lpage>983</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s11119-021-09870-3</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Johansen</surname> <given-names>K. L.</given-names>
</name>
<name>
<surname>Morton</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Malb&#xe9;teau</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Aragon</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Al-Mashharawi</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Ziliani</surname> <given-names>M. G.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>a). <article-title>Unmanned aerial vehicle-based phenotyping using morphometric and spectral analysis can quantify responses of wild tomato plants to salinity stress</article-title>. <source>Front. Plant Sci.</source> <volume>10</volume>, <elocation-id>370</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/fpls.2019.00370</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Johansen</surname> <given-names>K. L.</given-names>
</name>
<name>
<surname>Morton</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Malb&#xe9;teau</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Aragon</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Al-Mashharawi</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Ziliani</surname> <given-names>M. G.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>b). <article-title>Predicting biomass and yield at harvest of salt-stressed tomato plants using uav imagery</article-title>. <source>Int. Arch. Photogrammetry Remote Sens. Spatial Inf. Sci</source>. <page-range>407&#x2013;411</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.5194/ISPRS-ARCHIVES-XLII-2-W13-407-2019</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Johansen</surname> <given-names>K. L.</given-names>
</name>
<name>
<surname>Morton</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Malb&#xe9;teau</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Aragon</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Al-Mashharawi</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Ziliani</surname> <given-names>M. G.</given-names>
</name>
<etal/>
</person-group>. (<year>2020</year>). <article-title>Predicting biomass and yield in a tomato phenotyping experiment using UAV imagery and random forest</article-title>. <source>Front. Artif. Intell.</source> <volume>3</volume>, <elocation-id>28</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3389/frai.2020.00028</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Yao</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Omni-dimensional dynamic convolution</article-title>. <source>arXiv[Preprint]</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2209.07947</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Huang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<etal/>
</person-group>. (<year>2019</year>). <article-title>Effects of coupling of nursing seedling densities and seedling fetching area on transplanting quality and yield of hybrid rice</article-title>. <source>Trans. Chin. Soc Agric. Eng.</source> <volume>35</volume>, <fpage>20</fpage>&#x2013;<lpage>30</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.11975/j.issn.1002-6819.2019.24.003</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Yuan</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Research progress of rice transplanting mechanization</article-title>. <source>Trans. Chin. Soc Agric. Eng.</source> <volume>49</volume>, <fpage>1</fpage>&#x2013;<lpage>20</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.6041/j.issn.1000-1298.2018.05.001</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ma</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Tan</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Zeng</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Furrow opener for the precision drilling nursing seedlings of hybrid rice</article-title>. <source>Trans. Chin. Soc Agric. Eng.</source> <volume>39</volume>, <fpage>1</fpage>&#x2013;<lpage>14</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.11975/j.issn.1002-6819.202305140</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ma</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Yuan</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Hybrid rice achievements, development and prospect in China</article-title>. <source>J. Integr. Agric.</source> <volume>14</volume>, <fpage>197</fpage>&#x2013;<lpage>205</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/S2095-3119(14)60922-9</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nee</surname> <given-names>V. C.</given-names>
</name>
<name>
<surname>S.</surname> <given-names>C. L.</given-names>
</name>
<name>
<surname>Aijing</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Jianfeng</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Kitchen</surname> <given-names>N. R.</given-names>
</name>
<name>
<surname>Sudduth</surname> <given-names>K. A.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Corn emergence uniformity estimation and mapping using UAV imagery and deep learning</article-title>. <source>Comput. Electron. Agric.</source> <volume>198</volume>, <fpage>107008</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compag.2022.107008</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qi</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Seeding cavity detection in tray nursing seedlings of super rice based on computer vision technology</article-title>. <source>Trans. Chin. Soc Agric. Eng.</source> <volume>25</volume>, <fpage>121</fpage>&#x2013;<lpage>125</lpage>. doi:&#xa0;CNKI:SUN:NYGU.0.2009-02-025</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tan</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Liang</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Estimation on hole seeding quantity of super hybrid rice based on machine vision and BP neural network</article-title>. <source>Trans. Chin. Soc Agric. Eng.</source> <volume>30</volume>, <fpage>201</fpage>&#x2013;<lpage>208</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3969/j.issn.1002-6819.2014.21.024</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tong</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Wise-IoU: Bounding box regression loss with dynamic focusing mechanism</article-title>. <source>arXiv[Preprint]</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2301.10051</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Ding</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Research on the method of seeding quantity detection in potted seedling tray of super rice based on improved shape factor</article-title>. <source>Trans. Chin. Soc. Agric. Mach.</source> <volume>46</volume>, <fpage>29</fpage>&#x2013;<lpage>35 + 8</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.6041/j.issn.1000-1298.2015.11.005</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Bai</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Design and experiment on intelligent reseeding devices for rice tray nursing seedling based on machine vision</article-title>. <source>Trans. Chin. Soc Agric. Eng.</source> <volume>34</volume>, <fpage>35</fpage>&#x2013;<lpage>42</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.11975/j.issn.1002-6819.2018.13.005</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>Q.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Zuo</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Hu</surname> <given-names>Q.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>ECA-Net: Efficient channel attention for deep convolutional neural networks</article-title>,&#x201d; in <conf-name>2020 IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR).</conf-name> (<publisher-loc>Washington</publisher-loc>: <publisher-name>IEEE Press</publisher-name>), <fpage>11531</fpage>&#x2013;<lpage>11539</lpage>. doi: <pub-id pub-id-type="doi">10.1109/CVPR42600.2020.01155</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Ye</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Qian</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Continuous picking of yellow peaches with recognition and collision-free path</article-title>. <source>Comput. Electron. Agric.</source> <volume>214</volume>, <fpage>108273</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compag.2023.108273</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yang</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Bender</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Le</surname> <given-names>Q. V.</given-names>
</name>
<name>
<surname>Ngiam</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>CondConv: Conditionally parameterized convolutions for efficient inference</article-title>. <source>arXiv[Preprint]</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.1904.04971</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yuan</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>General situation of rice production and current status of hybrid rice development in Thailand</article-title>. <source>Hybrid Rice</source> <volume>33</volume>, <fpage>1</fpage>&#x2013;<lpage>2</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.16267/j.cnki.1005-3956.20180711.197</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yun</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Eichhorn</surname> <given-names>M. P.</given-names>
</name>
<etal/>
</person-group>. (<year>2024</year>). <article-title>Status, advancements and prospects of deep learning methods applied in forest studies</article-title>. <source>Int. J. Appl. Earth Observation Geoinfo.</source> <volume>131</volume>, <fpage>103938</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.jag.2024.103938</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Cheng</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Yan</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Wu</surname> <given-names>Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Detection of wheat scab fungus spores utilizing the Yolov5-ECA-ASFF network structure</article-title>. <source>Comput. Electron. Agric.</source> <volume>210</volume>, <fpage>107953</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compag.2023.107953</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhao</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Gao</surname> <given-names>B.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Performance detection system of tray precision seeder based on machine vision</article-title>. <source>Trans. Chin. Soc. Agric. Mach.</source> <volume>45</volume>, <fpage>24</fpage>&#x2013;<lpage>28</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.6041/j.issn.1000-1298.2014.S0.004</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Zheng</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Ye</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Ren</surname> <given-names>D.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>Distance-IoU loss: faster and better learning for bounding box regression</article-title>,&#x201d; in <conf-name>Proceedings of the AAAI Conference on Artificial Intelligence</conf-name> (<publisher-name>New York</publisher-name>: <publisher-loc>AAAI Press</publisher-loc>), Vol. <volume>34</volume>. <fpage>12993</fpage>&#x2013;<lpage>13000</lpage>. doi: <pub-id pub-id-type="doi">10.1609/aaai.v34i07.6999</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Intelligent resow cavity detection in tray nursing seedlings of super rice based on labview</article-title>. <source>J. Jiamusi University(Natural Sci. Edition)</source> <volume>30</volume>, <fpage>731</fpage>&#x2013;<lpage>733 + 736</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3969/j.issn.1008-1402.2012.05.024</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>