<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Plant Sci.</journal-id>
<journal-title>Frontiers in Plant Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Plant Sci.</abbrev-journal-title>
<issn pub-type="epub">1664-462X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpls.2025.1604514</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Plant Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Automatic detection of lucky bamboo nodes based on Improved YOLOv7</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Zhang</surname>
<given-names>Jing</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Deng</surname>
<given-names>Ruoling</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>*</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/3016926/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Cai</surname>
<given-names>Chengzhi</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/3064379/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Zou</surname>
<given-names>Erpeng</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liu</surname>
<given-names>Haitao</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2187626/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Hou</surname>
<given-names>Mingxin</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2586185/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Chen</surname>
<given-names>Xinzhi</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/3023485/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Lin</surname>
<given-names>Huamin</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Wei</surname>
<given-names>Zhenye</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>School of Mechanical Engineering, Guangdong Ocean University</institution>, <addr-line>Zhanjiang, Guangdong</addr-line>,&#xa0;<country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Guangdong Engineering Technology Research Center of Ocean Equipment and Manufacturing, Guangdong Ocean University</institution>, <addr-line>Zhanjiang, Guangdong</addr-line>,&#xa0;<country>China</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Guangdong Provincial Key Laboratory of Intelligent Equipment for South China Sea Marine Ranching, Guangdong Ocean University</institution>, <addr-line>Zhanjiang, Guangdong</addr-line>,&#xa0;<country>China</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>College of Mathematics and Computer Science, Guangdong Ocean University</institution>, <addr-line>Zhanjiang, Guangdong</addr-line>,&#xa0;<country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>Edited by: Chunlei Xia, Chinese Academy of Sciences (CAS), China</p>
</fn>
<fn fn-type="edited-by">
<p>Reviewed by: Xing Xu, South China Agricultural University, China</p>
<p>Lixiong Gong, Hubei University of Technology, China</p>
<p>Ankush Sawarkar, Shri Guru Gobind Singhji Institute of Engineering and Technology, India</p>
</fn>
<fn fn-type="corresp" id="fn001">
<p>*Correspondence: Ruoling Deng, <email xlink:href="mailto:dengruoling@gdou.edu.cn">dengruoling@gdou.edu.cn</email>
</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>17</day>
<month>07</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>16</volume>
<elocation-id>1604514</elocation-id>
<history>
<date date-type="received">
<day>02</day>
<month>04</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>30</day>
<month>06</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2025 Zhang, Deng, Cai, Zou, Liu, Hou, Chen, Lin and Wei</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Zhang, Deng, Cai, Zou, Liu, Hou, Chen, Lin and Wei</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<sec>
<title>Introduction</title>
<p>The detection of lucky bamboo (<italic>Dracaena sanderiana</italic>) nodes is a critical prerequisite for machining bamboo into high-value handicrafts. Current manual detection methods are inefficient, labor-intensive, and error-prone, necessitating an automated solution.</p>
</sec>
<sec>
<title>Methods</title>
<p>This study proposes an improved YOLOv7-based model for real-time, precise bamboo node detection. The model integrates a Squeeze-and-Excitation (SE) attention mechanism into the feature extraction network to enhance target localization and introduces a Weighted Intersection over Union (WIoU) loss function to optimize bounding box regression. A dataset of 2,000 annotated images (augmented from 1,000 originals) was constructed, covering diverse environmental conditions (e.g., blurred backgrounds, occlusions). Training was conducted on a server with an RTX 4090 GPU using PyTorch.</p>
</sec>
<sec>
<title>Results</title>
<p>The proposed model achieved a 97.6% mAP@0.5, significantly outperforming the original YOLOv7 (83.4% mAP) by 14.2%, while maintaining the same inference speed (100.18 FPS). Compared to state-of-the-art alternatives, our model demonstrated superior efficiency. It showed 41.5% and 153% higher FPS than YOLOv11 (70.8 FPS) and YOLOv12 (39.54 FPS), respectively. Despite marginally lower mAP (&#x2264;1.3%) versus these models, the balanced trade-off between accuracy and speed makes it more suitable for industrial deployment. Robustness tests under challenging conditions (e.g., low light, occlusions) further validated its reliability, with consistent confidence scores across scenarios.</p>
</sec>
<sec>
<title>Discussion</title>
<p>The proposed method significantly improves detection accuracy and efficiency, offering a viable tool for industrial applications in smart agriculture and handicraft production. Future work will address limitations in detecting nodes obscured by mottled patterns or severe occlusions by expanding label categories during training.</p>
</sec>
</abstract>
<kwd-group>
<kwd>lucky bamboo</kwd>
<kwd>handicraft</kwd>
<kwd>convolutional neural network</kwd>
<kwd>YOLOv7</kwd>
<kwd>object detection</kwd>
<kwd>bamboo node</kwd>
</kwd-group>
<counts>
<fig-count count="10"/>
<table-count count="3"/>
<equation-count count="10"/>
<ref-count count="35"/>
<page-count count="12"/>
<word-count count="5289"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-in-acceptance</meta-name>
<meta-value>Plant Bioinformatics</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1" sec-type="intro">
<label>1</label>
<title>Introduction</title>
<p>Lucky bamboo (<italic>Dracaena sanderiana</italic>) is a potted ornamental plant with excellent ornamental value (<xref ref-type="bibr" rid="B5">Chen, 2012</xref>; <xref ref-type="bibr" rid="B2">Akinlabi et&#xa0;al., 2017</xref>; <xref ref-type="bibr" rid="B11">Gergel and Turner, 2017</xref>). Traditional lucky bamboo is only used for flower arrangement or direct potting, with low added value (<xref ref-type="bibr" rid="B25">Rout et&#xa0;al., 2006</xref>; <xref ref-type="bibr" rid="B32">van Dam et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B1">Abdel-Rahman et&#xa0;al., 2020</xref>). Processing lucky bamboo into handicrafts can substantially increase its ornamental value, which is deeply loved by the public and has a large demand in the international market (<xref ref-type="bibr" rid="B28">Sharma et&#xa0;al., 2014</xref>; <xref ref-type="bibr" rid="B3">Amin and Mujeeb, 2019</xref>; <xref ref-type="bibr" rid="B19">Liu et&#xa0;al., 2019</xref>). Since the processing of lucky bamboo is cutting the lucky bamboo according to the bamboo nodes to meet different requirements, identifying the lucky bamboo nodes is the first and most crucial step in processing lucky bamboo. However, the existing methods of identifying lucky bamboo nodes still mainly rely on manual work, which has the disadvantages of low efficiency, high labour cost, and prone to errors. Therefore, studying a method that can automatically recognize lucky bamboo nodes with high efficiency and precision is imperative.</p>
<p>In recent years, traditional image processing methods have been widely used to identify bamboo-related fields. Juyal P et&#xa0;al. used methods such as logistic regression, support vector machine, naive Bayesian, random forest, convolutional neural network and ResNet to conduct a comparative analysis, and finally could accurately identify five common bamboos (<xref ref-type="bibr" rid="B15">Juyal et&#xa0;al., 2020</xref>). Watanabe used convolutional neural networks (CNN) to identify Japanese bamboo forest areas through Google satellite images, with an overall recognition accuracy of 93.7% (<xref ref-type="bibr" rid="B35">Watanabe et&#xa0;al., 2020</xref>). Kumar conducted research on bamboo leaf disease detection and developed a program to automatically recognize bamboo leaf diseases based on image processing and CNN (<xref ref-type="bibr" rid="B16">Kumar et&#xa0;al., 2022</xref>). Ziwei Wang used a residual neural network, original dataset, and MixUp dataset to optimize the traditional algorithm CNN to further improve the classification ability of bamboo species (<xref ref-type="bibr" rid="B34">Wang et&#xa0;al., 2022</xref>). Pankaja used Fourier descriptors to extract bamboo leaf features and used the Bayes classifier to identify bamboo leaves with an accuracy of 88.03% (<xref ref-type="bibr" rid="B21">Pankaja and Thippeswamy, 2017</xref>).</p>
<p>Existing target detection algorithms can be roughly divided into two categories. The first category is the two-stage R-CNN (<xref ref-type="bibr" rid="B12">He et&#xa0;al., 2017</xref>) series of algorithms based on region proposal, such as R-CNN, Fast R-CNN (<xref ref-type="bibr" rid="B17">Li et&#xa0;al., 2017</xref>), Faster R-CNN (<xref ref-type="bibr" rid="B26">Salvador et&#xa0;al., 2016</xref>), etc. The position of the object frame usually needs to be found first and then the category of the object frame will be determined by this algorithm. Although this type of method has high recognition accuracy, it takes a long time to calculate and is not suitable for real-time detection. The second category is the one-stage algorithm (<xref ref-type="bibr" rid="B30">Tian et&#xa0;al., 2022</xref>) represented by YOLO (<xref ref-type="bibr" rid="B22">Redmon, 2016</xref>) and SSD (<xref ref-type="bibr" rid="B18">Liu et&#xa0;al., 2016</xref>). This type of algorithm takes regression as the core, omitting the region proposal link of the two-stage algorithm, directly distinguishing specific categories and returning the bounding box (<xref ref-type="bibr" rid="B23">Redmon and Farhadi, 2017</xref>). Shilan Hong used optimized YOLOv4 to model the detection of bamboo shoots and proposed a classification and screening strategy to track each bamboo shoot (<xref ref-type="bibr" rid="B13">Hong et&#xa0;al., 2022</xref>). The experimental results showed that the average relative error and variance of the number of bamboo shoots were 1.28% and 0.016%, respectively, and the average relative error and variance of the corresponding pixel height results were -0.39% and 0.02%, respectively. The advantage of this method is that it can perform real-time detection, which is beneficial to improving the recognition efficiency of bamboo nodes. However, there is much room for improvement in the detection accuracy and robustness of this method, and its detection ability for small targets is also relatively poor. The current target detection algorithm cannot take into account both detection accuracy and timeliness, and the detection effect for small target units such as bamboo nodes is relatively poor.</p>
<p>To solve the above problems, this paper proposed a method for real-time and precise detection of luck bamboo nodes based on the improved model of YOLOv7. Firstly, an attention mechanism was introduced into the feature extraction network to enhance effective feature information and suppress invalid information. This can help the model locate and identify the lucky bamboo nodes faster and more accurately. Then, WIoU was introduced in the loss value calculation to optimize the bounding box regression process of the bamboo node through a dynamic weighting mechanism, thereby improving the model&#x2019;s detection ability for complex scenes and small bamboo nodes.</p>
<p>The main objectives of this research were to (a) construct the image dataset of lucky bamboo nodes using the data augmentation method, (b) establish the lucky bamboo node detection model using the improved YOLOv7, and (c) evaluate the detection stability and accuracy of the proposed method.</p>
</sec>
<sec id="s2" sec-type="materials|methods">
<label>2</label>
<title>Materials and methods</title>
<sec id="s2_1">
<label>2.1</label>
<title>Dataset construction</title>
<p>The image collecting equipment used in this study is shown in <xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref>. The equipment consisted of an USB industrial camera, an assembly line workbench, and a laptop. The assembly line workbench was equipped with conveyor belt, a conveyor belt motor, conveyor belt speed control controller, bamboo posture corrector, a cutting motor, a blade for cutting bamboo, a blade bearing housing, and conveyor belt base. The industrial camera was mounted on the tube, parallel to the conveyor belt.</p>
<fig id="f1" position="float">
<label>Figure&#xa0;1</label>
<caption>
<p>Acquisition equipment of lucky bamboo images.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g001.tif">
<alt-text content-type="machine-generated">Illustration of a bamboo cutting system, showing components like a laptop, conveyor belt with motor and speed control, industrial camera, bamboo posture corrector, and a cutting motor with blade for bamboo.</alt-text>
</graphic>
</fig>
<p>The lucky bamboo samples were planted in the processing base of Fugui Horticultural Farm in Mazhang District, Zhanjiang City, Guangdong Province, China. The images of lucky bamboo were collected using USB industrial camera connected to a laptop on November 1, 2023. The focal length of the camera lens is 3.6 mm. Specifically, the lucky bamboo was placed on the horizontal conveyor belt and then photographed at a vertical height of 50 mm from the lucky bamboo. Since the original lucky bamboo image had a resolution of 2 million pixels, which was too large and will increase the amount of calculation, OpenCV in Python was used to process the original image. Finally, the original images were compressed to 640&#xd7;640 pixels, with a horizontal and vertical resolutions of 96 dpi. Among them, the lucky bamboo image was stored in JPG format with a size of 32 MB. In the end, a total of 1,000 lucky bamboo images were collected, as shown in <xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>.</p>
<fig id="f2" position="float">
<label>Figure&#xa0;2</label>
<caption>
<p>Images of lucky bamboo.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g002.tif">
<alt-text content-type="machine-generated">Close-up of a green bamboo stem against a neutral background, showing the texture and segments of the plant.</alt-text>
</graphic>
</fig>
<p>To improve the generalization ability and robustness of the model and avoid model overfitting, the transforms module based on the Pytorch deep learning framework was used to perform data augmentation and expansion on the collected images (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3A</bold>
</xref>) (<xref ref-type="bibr" rid="B8">Cubuk et&#xa0;al., 2019</xref>). Since two major data enhancement methods (geometric transformation (flip) and color space perturbation (jitter + grayscale) had been widely proven to effectively improve the generalization ability of the model (<xref ref-type="bibr" rid="B29">Shorten and Khoshgoftaar, 2019</xref>). Therefore, the lucky bamboo images were performed random horizontal flip (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3B</bold>
</xref>), random vertical flip (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3C</bold>
</xref>), and color jitter (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3D</bold>
</xref>). Also, the RandomGrayscale function was used to convert the image to grayscale with a probability of 2.5% (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3E</bold>
</xref>), allowing the model to learn image features without color information and enhance the detection capabilities in complex scenes. After data enhancement, a total of 5000 images were obtained. Finally, 2000 representative images were obtained for constructing the lucky bamboo dataset. Also, the bamboo dataset was randomly divided into training set, verification set, and testing set in a ratio of 7:2:1 for model training and testing.</p>
<fig id="f3" position="float">
<label>Figure&#xa0;3</label>
<caption>
<p>Data augmentation of luck bamboo nodes. <bold>(A)</bold> original image; <bold>(B)</bold> random horizontal flip, <bold>(C)</bold> random vertical flip, <bold>(D)</bold> color jitter, <bold>(E)</bold> random grayscale.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g003.tif">
<alt-text content-type="machine-generated">Five images of a bamboo stalk section: (A), (B), and (C) show green bamboo from different angles on a light background. (D) presents the bamboo in dimmer lighting. (E) shows the same section in grayscale.</alt-text>
</graphic>
</fig>
<p>To enable the model to accurately locate bamboo nodes, the open-source software LabelImg was used to manually annotate all the lucky bamboo images after data augmentation (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4</bold>
</xref>) (<xref ref-type="bibr" rid="B9">Darrenl, 2017</xref>). In this study, the content of annotation was each node of lucky bamboo, that was, the collected location coordinate information. After labeling, all data were saved in the Pascal VOC dataset format. The annotation diagram is shown in <xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4</bold>
</xref>. Among them, the green frame represented the position of the bamboo nodes of lucky bamboo in the image.</p>
<fig id="f4" position="float">
<label>Figure&#xa0;4</label>
<caption>
<p>The image annotation of lucky bamboo nodes (<xref ref-type="bibr" rid="B9">Darrenl, 2017</xref>).</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g004.tif">
<alt-text content-type="machine-generated">A software interface displaying an image of a green bamboo node with rectangular boxes around sections of it. The interface includes tool options on the left and a list of labeled boxes on the right, with &#x201c;bamboo node&#x201d; checked.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s2_2">
<label>2.2</label>
<title>Bamboo node detection method</title>
<sec id="s2_2_1">
<label>2.2.1</label>
<title>Advantage of YOLOv7 model</title>
<p>YOLO (You Only Look Once) was a deep learning algorithm for real-time object detection that only required one forward propagation through a given neural network to detect all objects in the image (<xref ref-type="bibr" rid="B22">Redmon, 2016</xref>). This gave the YOLO algorithm an advantage over other algorithms in terms of speed, making it one of the most famous detection algorithms to date.</p>
<p>YOLO divided the image into multiple small grids and predicted multiple bounding boxes in each grid, as well as the object categories within each bounding box. YOLO used a single neural network to predict all bounding boxes and classes in an image, instead of using multiple neural networks to predict each bounding box. The advantage of YOLO was that it can run in real time and can detect more objects. YOLOv7 (<xref ref-type="bibr" rid="B33">Wang et&#xa0;al., 2023</xref>) was optimized on the basis of YOLOv4 (<xref ref-type="bibr" rid="B4">Bochkovskiy et&#xa0;al., 2020</xref>), training a better model with less training time. The input layer of YOLOv7 supported image enhancement (such as Mosaic) and adaptive anchor box calculation. Its backbone network used CSPDarknet53 and enhanced feature extraction capabilities by using CSPNet. Also, the Neck module in YOLOv7 combined Path Aggregation Network (PANet) and Spatial Pyramid Pooling (SPP) modules, which can better handle multi-scale features and enhance the accuracy of the model.</p>
<p>The algorithm mainly consisted of an input terminal, a feature extraction network, a feature fusion network, and an output terminal (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5</bold>
</xref>). Compared with YOLOv4, YOLOv7 employed a Focus operation that sampled the original image at double intervals in both horizontal and vertical directions. This reduced FLOPs value and computational complexity, thereby improving detection speed. Additionally, in the feature extraction module, YOLOv7 replaced the original CSP module with the C3 module. This modification enhanced training speed, reduced gradient redundancy, and improved learning efficiency. For network input processing, YOLOv3 (<xref ref-type="bibr" rid="B10">Farhadi and Redmon, 2018</xref>) and YOLOv4 required executing a separate program to compute initial anchor boxes when training on different datasets. In contrast, YOLOv7 integrated this functionality directly into the framework, enabling adaptive calculation of optimal anchor boxes for each training scenario. Also, unlike two-stage algorithms (e.g., Faster R-CNN), YOLOv7 eliminated the computationally intensive feature extraction and region proposal steps, significantly reducing inference time. Although the accuracy of YOLOv7 was slightly lower than that of Faster RCNN, its detection speed was faster and supported real-time detection. Therefore, YOLOv7 was selected as the basic framework for detecting lucky bamboo nodes.</p>
<fig id="f5" position="float">
<label>Figure&#xa0;5</label>
<caption>
<p>Architecture of the improved YOLOv7.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g005.tif">
<alt-text content-type="machine-generated">Flowchart of a C9_1 module incorporating an SE attention mechanism. The diagram shows processes including CBS, MP-C3, MaxPool, Concat, and Upsample. Each CBS includes Conv, BN, and Silu operations. Pathways connect various modules, leading to final operations including ADD, Conv, and Mul. The module redesign integrates the SE attention mechanism.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s2_2_2">
<label>2.2.2</label>
<title>Problems caused by using YOLOv7 model</title>
<p>YOLOv7 was currently the best-engineered generation object detection model of the YOLO series. It achieved real-time image processing speeds while maintaining accuracy compared to state-of-the-art models. Therefore, it was widely used in real-time vision applications. However, YOLOv7 still had some shortcomings in the application environment of lucky bamboo node detection in this paper, which were mainly reflected in the following aspects.</p>
<list list-type="simple">
<list-item>
<p>(a) The regression idea of the YOLO algorithm was to divide the image into S&#xd7;S grids, that was, each grid could only predict at most one object. Consequently, when multiple objects occupied the same grid, the algorithm&#x2019;s detection performance degraded significantly, often failing to identify all objects.</p>
</list-item>
<list-item>
<p>(b) The original network employed the Generalized Intersection over Union (GLoU) (<xref ref-type="bibr" rid="B24">Rezatofighi et&#xa0;al., 2019</xref>) loss function for bounding box regression. However, GIoU demonstrated limited effectiveness for small object detection, as it failed to explicitly incorporate object scale considerations. Furthermore, this loss function may induce a bias toward predicting larger bounding boxes, ultimately compromising detection accuracy.</p>
</list-item>
</list>
</sec>
<sec id="s2_2_3">
<label>2.2.3</label>
<title>Construction of bamboo node detection model</title>
<p>In response to the problems raised above, the following improvements were made to the YOLOv7 original network in this paper.</p>
<list list-type="simple">
<list-item>
<p>(a) Adding the SE attention mechanism (<xref ref-type="bibr" rid="B14">Hu et&#xa0;al., 2018</xref>; <xref ref-type="bibr" rid="B20">Niu et&#xa0;al., 2021</xref>) can make the network pay more attention to the bamboo nodes to be detected, thereby improving the model accuracy. The SE attention mechanism adaptively recalibrated channel-weight feature responses through its Squeeze and Excitation operations, selectively enhancing discriminative features. This mechanism enhanced feature representation accuracy by adaptively amplifying salient features while suppression irrelevant channel responses. Furthermore, the SE module enhanced model performance in complex scenarios through learned channel-weight adaptation, dynamically optimizing feature attention to improve both robustness and detection accuracy.</p>
</list-item>
<list-item>
<p>(b) To enhance model detection efficiency, we implemented the Weighted Intersection over Union (WIoU) bounding box loss function (<xref ref-type="bibr" rid="B7">Cho, 2021</xref>; <xref ref-type="bibr" rid="B31">Tong et&#xa0;al., 2023</xref>) at the network&#x2019;s output layer. The proposed method introduced an efficient IoU-based loss function that addressed the limitations of conventional approaches, achieving both accelerated convergence and enhanced regression accuracy. The WIoU loss function incorporated a dynamic non-monotonic focusing mechanism that more effectively evaluated anchor quality. This approach reduced dominance by high-quality anchors while mitigating harmful gradients from low-quality samples. This allowed the WIoU loss function to focus on anchor frames of ordinary quality and improve overall detection performance.</p>
</list-item>
</list>
<sec id="s2_2_3_1">
<label>2.2.3.1</label>
<title>SE attention mechanism</title>
<p>In the traditional convolutional neural network (CNN) architecture, convolutional layers and pooling layers were the core components for building deep feature representations (<xref ref-type="bibr" rid="B27">Sawarkar et&#xa0;al., 2024</xref>). These layers formed hierarchical feature representations by gradually extracting local features in the image and reducing the spatial dimension of the data. However, there was an implicit assumption in this process. That was, each channel of the feature map was equally important to the final task. Nevertheless, in practical applications, different channels often carry different amounts of information or importance, and contribute differently to the task (<xref ref-type="bibr" rid="B6">Chen et&#xa0;al., 2017</xref>). To solve this problem, Hu et&#xa0;al. proposed the SE attention mechanism architecture, which improved the model&#x2019;s ability to express features by adaptively recalibrating the importance of each channel (<xref ref-type="bibr" rid="B14">Hu et&#xa0;al., 2018</xref>). Firstly, global average pooling was used to capture the global information of each channel, and then a channel weight vector was generated using a fully connected layer and a sigmoid activation function. Then, this weight vector was applied to each channel of the original feature map, and the feature map was scaled by element-by-element multiplication between channels, thereby enhancing the feature representation of those important channels and weakening the influence of those irrelevant or redundant channels. This adaptive channel recalibration mechanism enabled the model to focus more on the features that contributed most to the task, thereby improving the performance and generalization ability of the network. Therefore, a three-layer SE attention mechanism was added to the backbone network of the original YOLOv7 model. The structure diagram of the SE attention mechanism is shown in <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>.</p>
<fig id="f6" position="float">
<label>Figure&#xa0;6</label>
<caption>
<p>The structure of SE attention mechanism.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g006.tif">
<alt-text content-type="machine-generated">Diagram of a transformation process from input tensor \( X \) to output tensor \( \dot{X} \) using functions \( F_{tr} \), \( F_{sq} \), \( F_{ex} \), and \( F_{scale} \). Initial tensor \( X \) with dimensions \( C', H', W' \) is transformed, analyzed, and scaled, resulting in output with dimensions \( C, H, W \). Arrows indicate the sequence of operations.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s2_2_3_2">
<label>2.2.3.2</label>
<title>WIoU loss function</title>
<p>The Generalized Intersection over Union (GIoU) loss function was used in the original YOLOv7 architecture (<xref ref-type="bibr" rid="B4">Bochkovskiy et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B33">Wang et&#xa0;al., 2023</xref>). In most cases, GIOU can calculate IoU at a wide level. When predicted boxes perfectly coincided with ground truth boxes, their intersection area equaled their individual areas. Also, their minimum bounding rectangles were also the same. In this case, the GIoU value saturated at 1, making it unable to distinguish subtle deviations in predicted box alignment. Under this condition, GIoU reduced to standard IoU. Furthermore, the GIoU loss function suffered from two key limitations: (a) higher computational overhead, and (b) slower convergence compared to more recent alternatives. To calculate GIoU, it was necessary to find the minimum bounding rectangle for each predicted box and the true box. This approach introduced significant computational overhead and adversely impacted training convergence, particularly when processing high-volume datasets or high-resolution imagery. To solve this problem, this paper adopted the Weighted Intersection over Union (WIoU) loss function in the improved network. WIoU incorporated a dynamic non-monotonic mechanism for bounding box regression that adaptively modulated gradient distributions based on overlap states, effectively mitigating both excessive and harmful gradients form outlier samples. Through optimized gradient allocation, the WIoU loss function enhanced model performance in normal cases while demonstrating superior robustness in extreme scenarios, simultaneously accelerating convergence and improving training efficiency. Consequently, this study replaced the original YOLOv7&#x2019;s GIoU loss function with the WIoU variant to enhance bounding box regression performance. The calculation process of GIoU and WIoU loss function were shown in the following <xref ref-type="disp-formula" rid="eq1">Equations 1</xref>&#x2013;<xref ref-type="disp-formula" rid="eq6">6</xref>.</p>
<disp-formula id="eq1">
<label>(1)</label>
<mml:math display="block" id="M1">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mtext>o</mml:mtext>
<mml:mi>U</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mo>|</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mi>p</mml:mi>
</mml:msub>
<mml:mo>&#x2229;</mml:mo>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>|</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mo>|</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mi>p</mml:mi>
</mml:msub>
<mml:mo>&#x222a;</mml:mo>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>|</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq2">
<label>(2)</label>
<mml:math display="block" id="M2">
<mml:mrow>
<mml:mi>G</mml:mi>
<mml:mi>I</mml:mi>
<mml:mtext>o</mml:mtext>
<mml:mi>U</mml:mi>
<mml:mo>=</mml:mo>
<mml:mi>I</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>U</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mo>|</mml:mo>
<mml:mrow>
<mml:mi>C</mml:mi>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mo>\</mml:mo>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mi>P</mml:mi>
</mml:msub>
<mml:mo>&#x222a;</mml:mo>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo>|</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mo>|</mml:mo>
<mml:mi>C</mml:mi>
<mml:mo>|</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq3">
<label>(3)</label>
<mml:math display="block" id="M3">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mi>G</mml:mi>
<mml:mi>I</mml:mi>
<mml:mtext>o</mml:mtext>
<mml:mi>U</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>G</mml:mi>
<mml:mi>I</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>U</mml:mi>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq4">
<label>(4)</label>
<mml:math display="block" id="M4">
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mi>&#x3bd;</mml:mi>
<mml:mi>&#x3b1;</mml:mi>
</mml:mfrac>
<mml:mo>.</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mo>|</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mtext>p</mml:mtext>
</mml:msub>
<mml:mo>&#x2229;</mml:mo>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>|</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mo>|</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mtext>p</mml:mtext>
</mml:msub>
<mml:mo>&#x222a;</mml:mo>
<mml:msub>
<mml:mi>B</mml:mi>
<mml:mrow>
<mml:mi>g</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>|</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq5">
<label>(5)</label>
<mml:math display="block" id="M5">
<mml:mrow>
<mml:mi>W</mml:mi>
<mml:mi>I</mml:mi>
<mml:mtext>o</mml:mtext>
<mml:mi>U</mml:mi>
<mml:mo>=</mml:mo>
<mml:mi>I</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>U</mml:mi>
<mml:mo>&#xb7;</mml:mo>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq6">
<label>(6)</label>
<mml:math display="block" id="M6">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mi>W</mml:mi>
<mml:mi>I</mml:mi>
<mml:mtext>o</mml:mtext>
<mml:mi>U</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>W</mml:mi>
<mml:mi>I</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>U</mml:mi>
</mml:mrow>
</mml:math>
</disp-formula>
<p>Where <italic>B<sub>p</sub>
</italic> is the predicted bounding box, <italic>B<sub>gt</sub>
</italic> is the true bounding box, <italic>C</italic> is the smallest rectangle that can contain both the predicted box <italic>B<sub>p</sub>
</italic> and the true box <italic>B<sub>gt</sub>
</italic>, <italic>&#x3b2;</italic> is dynamic weight, <italic>&#x3bd;</italic> is a normalization factor and <italic>&#x3b1;</italic> is a learnable parameter.</p>
</sec>
</sec>
</sec>
<sec id="s2_3">
<label>2.3</label>
<title>Experimental environment</title>
<p>The experiment platform in this research was an autonomously configured server running in the deep learning framework. The hardware environment is Intel Core i9-14900KF processor with 64GB running memory, and the graphics card is Nvidia GeForce RTX 4090 with 24GB memory. The software environment is a virtual environment built using Anaconda under the Ubuntu20.04 operating system. The virtual environment consisted of Pytorch 2.1.0, CUDA 12.3, and Python 3.8.6. The Python language was used as the main language for writing program codes. Also, the numpy, pandas, OpenCV and other required libraries were called to implement the training and testing of the lucky bamboo node detection model. The hardware and software configurations for the established model were listed in <xref ref-type="table" rid="T1">
<bold>Table 1</bold>
</xref>.</p>
<table-wrap id="T1" position="float">
<label>Table&#xa0;1</label>
<caption>
<p>The hardware and software configurations for the established model.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="center">Project</th>
<th valign="top" align="center">Content</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">Operating system</td>
<td valign="top" align="center">Ubuntu20.04</td>
</tr>
<tr>
<td valign="top" align="center">CPU</td>
<td valign="top" align="center">Intel Core i9-14900KF</td>
</tr>
<tr>
<td valign="top" align="center">GPU</td>
<td valign="top" align="center">NVIDIA GeForce RTX 4090</td>
</tr>
<tr>
<td valign="top" align="center">RAM</td>
<td valign="top" align="center">64GB</td>
</tr>
<tr>
<td valign="top" align="center">Compiled language</td>
<td valign="top" align="center">Python3.8.6</td>
</tr>
<tr>
<td valign="top" align="center">Deep learning framework</td>
<td valign="top" align="center">CUDA12.3 Pytorch 2.1.0</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s2_4">
<label>2.4</label>
<title>Evaluation of bamboo node detection model</title>
<p>In this study, objective indicators such as precision, recall, and loss function convergence curve would be used to evaluate the performance of bamboo node detection model. Among them, Intersection over Union (IOU) represents the ratio of the intersection and union between the detected bounding box and the real bounding box, which is a common indicator for evaluating the performance of object detection model. The higher the IOU value, the better the model detection performance. Precision refers to the ratio of the number of correctly detected bamboo nodes to the total number of detected bamboo nodes. Recall is the ratio of the number of correctly detected bamboo nodes to the number of actual true bamboo nodes. The precision and recall were computed by <xref ref-type="disp-formula" rid="eq7">Equations 7</xref> and <xref ref-type="disp-formula" rid="eq8">8</xref>. True Positive (<italic>TP</italic>) indicates the number of correctly detected lucky bamboo nodes when the IOU is greater than or equal to the selected threshold. False Positive (FP) indicates the number of misjudged bamboo nodes when the IOU is smaller than the selected threshold. False Negative (<italic>FN</italic>) represents the number of undetected bamboo nodes. Average Precision (<italic>AP</italic>) refers to the area under the Precision-Recall curve. The higher the <italic>AP</italic> value, the better the performance of the model in detecting bamboo nodes. Mean Average Precision (<italic>mAP</italic>) refers to the mean value of <italic>AP</italic> for all categories. The AP and <italic>mAP</italic> were computed by <xref ref-type="disp-formula" rid="eq9">Equations 9</xref> and <xref ref-type="disp-formula" rid="eq10">10</xref>.</p>
<disp-formula id="eq7">
<label>(7)</label>
<mml:math display="block" id="M7">
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>+</mml:mo>
<mml:mi>F</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq8">
<label>(8)</label>
<mml:math display="block" id="M8">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>+</mml:mo>
<mml:mi>F</mml:mi>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq9">
<label>(9)</label>
<mml:math display="block" id="M9">
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>=</mml:mo>
<mml:mstyle displaystyle="true">
<mml:mrow>
<mml:msubsup>
<mml:mo>&#x222b;</mml:mo>
<mml:mn>0</mml:mn>
<mml:mi>1</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mi>R</mml:mi>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:mstyle>
<mml:mi>d</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="eq10">
<label>(10)</label>
<mml:math display="block" id="M10">
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:msub>
<mml:mi>P</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:math>
</disp-formula>
</sec>
</sec>
<sec id="s3" sec-type="results">
<label>3</label>
<title>Results and discussion</title>
<sec id="s3_1">
<label>3.1</label>
<title>Performance of lucky bamboo node detection model</title>
<p>The training results of lucky bamboo node detection model are shown in <xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref>. With the increase of training epochs, the precision value, and mAP value of lucky bamboo node detection model gradually increased. The details are as follows. During epochs 0-20, the precision values of the model increased rapidly. Afterwards, the precision values of the model remained stable between 0.95 and 0.97. Similarly, the mAP of the model increased rapidly during epochs 0&#x2013;20 and then remained stable between 0.96 and 0.99. Conclusively, the training was suspended at 300 epochs.</p>
<fig id="f7" position="float">
<label>Figure&#xa0;7</label>
<caption>
<p>Comparison of training results between the original model and the lucky bamboo node detection model. <bold>(A)</bold> precision curve, <bold>(B)</bold> mAP curve.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g007.tif">
<alt-text content-type="machine-generated">Two line graphs compare metrics over epochs. Graph (A) shows Precision, and Graph (B) shows mAP@0.5. Both graphs have two lines: an orange line for &#x201c;sc_Wiou&#x201d; and a blue line for &#x201c;original&#x201d;. The orange line consistently outperforms the blue line, stabilizing close to 1, indicating higher performance.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3_2">
<label>3.2</label>
<title>Detection result of lucky bamboo node detection model</title>
<p>The lucky bamboo node detection model was tested on additional 400 RGB images collected later. <xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8</bold>
</xref> shows some examples of the detected bamboo nodes in different environmental conditions. The confidence scores were indicated beside each detected node. The high scores (up to 1.000) demonstrated that the results were quite reliable. In <xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8A</bold>
</xref>, even though the lucky bamboo plant was in low-light environment and blurry, the bamboo nodes thereon can still be correctly detected. In images taken in high-light conditions, the bamboo nodes appeared brighter and the bamboo nodes were similar in color to the background and surrounded by black shadows. It would be hard to recognize them manually, but the developed model was able to detect all the nodes in the image with high confident scores (<xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8B</bold>
</xref>). In <xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8C</bold>
</xref>, the bamboo nodes were occluded by a bamboo leaf, and one of them was occluded by more than 50%. But they were all successfully detected. In images taken at a high shooting distance condition (<xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8D</bold>
</xref>), the bamboo nodes accounted for a small proportion of pixels in the image, which would be difficult to be recognized. Nevertheless, every bamboo node was recognized with a high confidence score.</p>
<fig id="f8" position="float">
<label>Figure&#xa0;8</label>
<caption>
<p>Detection results of lucky bamboo node detection model. <bold>(A)</bold> low light condition, <bold>(B)</bold> high light condition, <bold>(C)</bold> complex condition, <bold>(D)</bold> 10 cm shooting distance.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g008.tif">
<alt-text content-type="machine-generated">Four close-up images labeled A, B, C, and D show sections of bamboo with labeled confidence scores displayed in colored boxes. Each score represents the likelihood of the identified bamboo segments, with varying levels of confidence indicated on the bamboo pieces. A white background is visible in all images.</alt-text>
</graphic>
</fig>
<p>Among the 400 images, a total of 1385 bamboo nodes were manually counted from the captured images by two people. With the same 400 images, the lucky bamboo node detection model detected 1376 true positives, 84 false positives, and 9 false negatives. Comparing two sets of results, the lucky bamboo node detection model was in good agreement with the manual counting, as indicated by the low errors, for example, 0.346 for RMSE. The minor discrepancies between the lucky bamboo node detection model and the manual counting could attribute to the following reasons. When training the CNN model, the white mottled rings of lucky bamboo (<xref ref-type="fig" rid="f9">
<bold>Figure&#xa0;9A</bold>
</xref>), the blocky mottled patches (<xref ref-type="fig" rid="f9">
<bold>Figure&#xa0;9B</bold>
</xref>), and the blurred bamboo body with scratches (<xref ref-type="fig" rid="f9">
<bold>Figure&#xa0;9C</bold>
</xref>) were not labeled when manually annotating the bamboo nodes. This meant that the model was not trained for these cases. This may lead to false positives (FP) in the detection. Avoiding these cases would improve the accuracy of bamboo node detection. In some cases, dry bamboo leaves that were not completely removed severely blocked the bamboo nodes, causing the CNN model to mistakenly classify the nearby nodes as background (<xref ref-type="fig" rid="f9">
<bold>Figure&#xa0;9D</bold>
</xref>). All these would explain the higher nodes counts of the bamboo node detection model. The resulting discrepancies could be minimized by labeling white mottled rings, blocky mottled patches, and blurred bamboo body with scratches as additional categories during image preprocessing, which would be explored in future studies. Overall, the low error showed that the model could successfully detect most of the bamboo nodes under various environmental conditions.</p>
<fig id="f9" position="float">
<label>Figure&#xa0;9</label>
<caption>
<p>Examples of images which might have caused bamboo nodes detection errors. <bold>(A)</bold> blurred conditions, <bold>(B)</bold> similarities of bamboo spot and node, <bold>(C)</bold> sick bamboo node, <bold>(D)</bold> dried bamboo leaf.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g009.tif">
<alt-text content-type="machine-generated">Close-up images of green stems with various conditions. (A) Healthy stem with a smooth surface. (B) Stem with a rough, damaged area. (C) Blurred green stem. (D) Stem showing peeling or frayed section.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3_3">
<label>3.3</label>
<title>Ablation experiment</title>
<p>To further verify that the bamboo node detection model with added SE attention mechanism and WIoU loss function performs better than the original YOLOv7 model, four ablation experiments were conducted. The same data set and the same training parameters and methods were used to complete the training of each set of experiments. The experimental results are shown in <xref ref-type="table" rid="T2">
<bold>Table&#xa0;2</bold>
</xref>, where &#x201c;&#xd7;&#x201d; represents that the corresponding improvement strategy is not used in the network model, and &#x201c;&#x221a;&#x201d; represents that the improvement strategy is used. Among them, Model 1 is the original YOLOv7 model, and model 2, model 3 and model 4 are models that add the SE attention mechanism, WIoU loss function, and SE attention mechanism-WIoU loss function to the original YOLOv7 model, respectively. The results showed that the detection accuracy (mAP) of the model 1 without any improvement strategy was 83.4%. Model 2, which introduced the SE attention mechanism based on the original YOLOv7 model, embedded spatial position information into channel attention. This enabled the model to achieve better prediction results when detecting bamboo nodes that relied on position information, thereby improving the detection accuracy (mAP) by 15.5%. Model 3 introduced a new bounding box loss function MIoU to reduce the overlap error between the predicted box and the true box, thereby improving the detection accuracy and stability of the model. Therefore, the detection accuracy (mAP) of model3 was improved by 12.7%. Model 4 (i.e. our developed model), which introduced the SE attention mechanism and MIoU loss function based on the original model, optimized the prediction accuracy of the YOLOv7 model for the recognition and localization of bamboo nodes, thereby improving the model detection accuracy (mAP) by 14.2%. Over all the four pretrained models, the model 4 had the best mAP (97.6%).</p>
<table-wrap id="T2" position="float">
<label>Table&#xa0;2</label>
<caption>
<p>Results of ablation experiment.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" rowspan="2" align="center">Model name</th>
<th valign="top" colspan="2" align="center">Improvement strategies</th>
<th valign="middle" rowspan="2" align="center">mAP (%)</th>
</tr>
<tr>
<th valign="top" align="center">SE-mechanism</th>
<th valign="top" align="center">WIoU</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">Model 1</td>
<td valign="top" align="center">&#xd7;</td>
<td valign="top" align="center">&#xd7;</td>
<td valign="top" align="center">83.4</td>
</tr>
<tr>
<td valign="top" align="center">Model 2</td>
<td valign="top" align="center">&#x2713;</td>
<td valign="top" align="center">&#xd7;</td>
<td valign="top" align="center">98.9</td>
</tr>
<tr>
<td valign="top" align="center">Model 3</td>
<td valign="top" align="center">&#xd7;</td>
<td valign="top" align="center">&#x2713;</td>
<td valign="top" align="center">96.1</td>
</tr>
<tr>
<td valign="top" align="center">Proposed model</td>
<td valign="top" align="center">&#x2713;</td>
<td valign="top" align="center">&#x2713;</td>
<td valign="top" align="center">97.6</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>To further evaluate the improved models, the P-R curves for the four models were plotted. The P-R curves showed that the model 4 had the highest precision consistently over the recall range of 0.90 to 1.00 (<xref ref-type="fig" rid="f10">
<bold>Figure&#xa0;10</bold>
</xref>). All these performance indicators proved that our proposed lucky bamboo node detection model was superior than the other three models.</p>
<fig id="f10" position="float">
<label>Figure&#xa0;10</label>
<caption>
<p>P-R curve of ablation experiment.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1604514-g010.tif">
<alt-text content-type="machine-generated">Precision-recall curve showing four models: Model 1 in blue, Model 2 in green, Model 3 in orange, and a proposed model in red. The proposed model outperforms others, maintaining higher precision across various recall levels. An inset zooms in on the high-recall region, highlighting detailed performance differences.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3_4">
<label>3.4</label>
<title>Comparison with other latest models</title>
<p>To better demonstrate the excellent performance of the proposed model, the proposed model was compared with the current mainstream YOLOv11 and YOLOv12 models. Compared to the baseline YOLOv7, our proposed model demonstrated comparable inference speed but achieved a significant 17.03% improvement in mAP (0.976). When benchmarked against Model 2, YOLOv11, and YOLOv12, our proposed model showed moderate mAP reductions of 1.3%, 1.22%, and 0.41% respectively. However, it delivered substantial Frame Per Second (FPS) enhancements-increasing FPS from 75.20 to 100.18 against Model 2, 70.8 to 100.18 versus YOLOv11, and 39.54 to 100.18 compared to YOLOv12 (<xref ref-type="table" rid="T3">
<bold>Table&#xa0;3</bold>
</xref>). These computational efficiency gains represent critical advantages for industrial deployment scenarios. Consequently, our proposed model exhibited reduced inference time and higher processing throughput, demonstrating particular advantages for video stream analysis and enhanced suitability for real-world deployment scenarios.</p>
<table-wrap id="T3" position="float">
<label>Table&#xa0;3</label>
<caption>
<p>Comparison results with other latest models.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="center">Model name</th>
<th valign="top" align="center">Parameter quantity (MB)</th>
<th valign="top" align="center">Average inference time (ms)</th>
<th valign="top" align="center">Frames Per Second (FPS)</th>
<th valign="top" align="center">mAP@0.5 (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="center">
<bold>Proposed model</bold>
</td>
<td valign="middle" align="center">136.04</td>
<td valign="middle" align="center">
<bold>9.98</bold>
</td>
<td valign="middle" align="center">
<bold>100.18</bold>
</td>
<td valign="middle" align="center">0.976</td>
</tr>
<tr>
<td valign="top" align="center">YOLOv7_original</td>
<td valign="middle" align="center">135.52</td>
<td valign="middle" align="center">9.98</td>
<td valign="middle" align="center">100.18</td>
<td valign="middle" align="center">0.834</td>
</tr>
<tr>
<td valign="top" align="center">Model 2</td>
<td valign="middle" align="center">135.52</td>
<td valign="middle" align="center">13.30</td>
<td valign="middle" align="center">75.20</td>
<td valign="middle" align="center">
<bold>0.989</bold>
</td>
</tr>
<tr>
<td valign="top" align="center">YOLOv11</td>
<td valign="middle" align="center">
<bold>109.11</bold>
</td>
<td valign="middle" align="center">14.12</td>
<td valign="middle" align="center">70.8</td>
<td valign="middle" align="center">0.988</td>
</tr>
<tr>
<td valign="top" align="center">YOLOv12</td>
<td valign="middle" align="center">227.44</td>
<td valign="middle" align="center">25.29</td>
<td valign="middle" align="center">39.54</td>
<td valign="middle" align="center">0.980</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>Bold values indicate the best performance achieved in each column (i.e., lowest parameter quantity, fastest inference time, highest FPS, and highest mAP@0.5).</p>
</fn>
</table-wrap-foot>
</table-wrap>
</sec>
</sec>
<sec id="s4" sec-type="conclusions">
<label>4</label>
<title>Conclusion</title>
<p>In this study, a high-precision and high-efficiency method was proposed based on a deep learning CNN mode for automatic detection of lucky bamboo node on the bamboo plant. Using the method, a high-throughput and low-cost application was developed and evaluated using lucky bamboo plant samples. The following conclusions were drawn. The CNN-based lucky bamboo node detection model was capable of recognizing and locating bamboo node in the lucky bamboo plant. The model was found to be the most efficient model for recognizing and locating bamboo nodes on the lucky bamboo plant structure. When compared to manual detection of bamboo node, the developed method had an estimated accuracy of 97.6%. The accuracy of the developed method was not affected by complex environment. The developed method shows great promise as a robust tool for computer-aided detecting of bamboo nodes on the lucky bamboo plant structure, which will help artisans rapidly and accurately identify lucky bamboo nodes to speed up the processing of lucky bamboo. However, more tests may be required to further verify the developed method.</p>
</sec>
</body>
<back>
<sec id="s5" sec-type="data-availability">
<title>Data availability statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec id="s6" sec-type="author-contributions">
<title>Author contributions</title>
<p>JZ: Conceptualization, Funding acquisition, Methodology, Project administration, Resources, Software, Supervision, Writing &#x2013; original draft, Writing &#x2013; review &amp; editing. RD: Data curation, Formal Analysis, Software, Validation, Writing &#x2013; original draft, Conceptualization, Funding acquisition, Methodology, Project administration, Supervision, Writing &#x2013; review &amp; editing. CC: Formal Analysis, Software, Visualization, Writing &#x2013; original draft. EZ: Data curation, Formal Analysis, Software, Writing &#x2013; original draft, Visualization. HaL: Data curation, Formal Analysis, Funding acquisition, Writing &#x2013; review &amp; editing. MH: Data curation, Funding acquisition, Writing &#x2013; review &amp; editing. XC: Data curation, Validation, Writing &#x2013; original draft. HuL: Writing &#x2013; review &amp; editing, Validation. ZW: Validation, Writing &#x2013; review &amp; editing.</p>
</sec>
<sec id="s7" sec-type="funding-information">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research and/or publication of this article. This research was funded by Guangdong Basic and Applied Basic Research Foundation (No. 2022A1515110468), the Science and Technology Projects of Zhanjiang (No. 2024B01044), the Innovation Team Project for Ordinary Universities in Guangdong Province (No. 2024KCXTD041), Program for scientific research startup funds of Guangdong Ocean University (No. 060302062106 and 060302062202), Natural Science Foundation of Guangdong Province (No. 2025A1515012901), and Guangdong Provincial Science and Technology Special Fund (No. A24333), Zhanjiang Key Laboratory of Modern Marine Fishery Equipment (No. 2021A05023).</p>
</sec>
<ack>
<title>Acknowledgments</title>
<p>The authors thank their partners, Zhanjiang Key Laboratory of Modern Marine Fishery Equipment, for providing the platform for research.</p>
</ack>
<sec id="s8" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s9" sec-type="ai-statement">
<title>Generative AI statement</title>
<p>The author(s) declare that no Generative AI was used in the creation of this manuscript.</p>
</sec>
<sec id="s10" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Abdel-Rahman</surname> <given-names>T. F.</given-names>
</name>
<name>
<surname>El-Morsy</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Halawa</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Occurrence of stem and leaf spots on lucky bamboo (Dracaena sanderiana hort. ex. mast.) plants in vase and its cure with safe means</article-title>. <source>J. Plant Prot. Pathol.</source> <volume>11</volume>, <fpage>705</fpage>&#x2013;<lpage>713</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.21608/jppp.2020.170648</pub-id>
</citation></ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Akinlabi</surname> <given-names>E. T.</given-names>
</name>
<name>
<surname>Anane-Fenin</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Akwada</surname> <given-names>D. R.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Bamboo</article-title>. <source>Multipurpose. Plant</source> <volume>268</volume>, <page-range>262</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/978-3-319-56808-9</pub-id>
</citation></ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Amin</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Mujeeb</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Callus induction and synthetic seed development in Draceana sanderiana Sanderex Mast: Lucky Bamboo</article-title>. <source>Biotechnol. J. Int.</source> <volume>23</volume>, <fpage>1</fpage>&#x2013;<lpage>8</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.9734/bji/2019/v23i330082</pub-id>
</citation></ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bochkovskiy</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>C.-Y.</given-names>
</name>
<name>
<surname>Liao</surname> <given-names>H.-Y. M.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Yolov4: Optimal speed and accuracy of object detection</article-title>. <source>arXiv. preprint. arXiv:.10934</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2004.10934</pub-id>
</citation></ref>
<ref id="B5">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2012</year>). <source>Landscape architecture: planting design illustrated</source> (<publisher-loc>Virginia, USA</publisher-loc>: <publisher-name>ArchiteG, Inc</publisher-name>).</citation></ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Xiao</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Jin</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Yan</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Feng</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Dual path networks</article-title>. <source>Adv. Neural Inf. Process. Syst.</source> <volume>30</volume>, <fpage>1</fpage>&#x2013;<lpage>11</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.1707.01629</pub-id>
</citation></ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cho</surname> <given-names>Y.-J.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Weighted intersection over union (wIoU): a new evaluation metric for image segmentation</article-title>. <source>Pattern Recognition Letters</source>. <volume>185</volume>, <page-range>101&#x2013;107</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.patrec.2024.07.011</pub-id>
</citation></ref>
<ref id="B8">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Cubuk</surname> <given-names>E. D.</given-names>
</name>
<name>
<surname>Zoph</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Mane</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Vasudevan</surname> <given-names>V.</given-names>
</name>
<name>
<surname>Le</surname> <given-names>Q. V.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>Autoaugment: Learning augmentation strategies from data</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</conf-name>. (<publisher-loc>California, USA</publisher-loc>). <fpage>113</fpage>&#x2013;<lpage>123</lpage>.</citation></ref>
<ref id="B9">
<citation citation-type="web">
<person-group person-group-type="author">
<collab>Darrenl</collab>
</person-group> (<year>2017</year>). <article-title>Labelimg: Labelimg Is a Graphical Image Annotation Tool and Label Object Bounding Boxes in Images</article-title>. Available online at: <uri xlink:href="https://github.com/tzutalin/labelImg">https://github.com/tzutalin/labelImg</uri> (Accessed <access-date>21 January 2025</access-date>).</citation></ref>
<ref id="B10">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Farhadi</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Redmon</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Yolov3: An incremental improvement</article-title>,&#x201d; in <source>Computer vision and pattern recognition</source> (<publisher-name>Springer Berlin</publisher-name>, <publisher-loc>Heidelberg, Germany</publisher-loc>), <fpage>1</fpage>&#x2013;<lpage>6</lpage>.</citation></ref>
<ref id="B11">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Gergel</surname> <given-names>S. E.</given-names>
</name>
<name>
<surname>Turner</surname> <given-names>M. G.</given-names>
</name>
</person-group> (<year>2017</year>). <source>Learning landscape ecology: a practical guide to concepts and techniques</source> (<publisher-loc>New York, USA</publisher-loc>: <publisher-name>Springer</publisher-name>).</citation></ref>
<ref id="B12">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>He</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Gkioxari</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Doll&#xe1;r</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Girshick</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>Mask r-cnn</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE international conference on computer vision</conf-name>. (<publisher-loc>Venice, Italy</publisher-loc>) <fpage>2961</fpage>&#x2013;<lpage>2969</lpage>.</citation></ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hong</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Jiang</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Zhu</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Rao</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Gao</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>A deep learning-based system for monitoring the number and height growth rates of moso bamboo shoots</article-title>. <source>Appl. Sci.</source> <volume>12</volume>, <fpage>7389</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/app12157389</pub-id>
</citation></ref>
<ref id="B14">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Hu</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Shen</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Sun</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Squeeze-and-excitation networks</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE conference on computer vision and pattern recognition</conf-name>. (<publisher-loc>Utah, USA</publisher-loc>) <fpage>7132</fpage>&#x2013;<lpage>7141</lpage>.</citation></ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Juyal</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Kulshrestha</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Sharma</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Ghanshala</surname> <given-names>T.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Common bamboo species identification using machine learning and deep learning algorithms</article-title>. <source>Int. J. Innovative. Technol. Explor. Eng.</source> <volume>9</volume>, <page-range>3012&#x2013;3017</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.35940/ijitee.D1609.029420</pub-id>
</citation></ref>
<ref id="B16">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Kumar</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Sharma</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Pandey</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Goyal</surname> <given-names>H. R.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>Identification of various bamboo diseases using deep learning approach</article-title>,&#x201d; in <conf-name>2022 IEEE Conference on Interdisciplinary Approaches in Technology and Management for Social Innovation (IATMSI)</conf-name>. <fpage>1</fpage>&#x2013;<lpage>6</lpage> (<publisher-loc>Gwalior, India</publisher-loc>: <publisher-name>IEEE</publisher-name>).</citation></ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Liang</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Shen</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Feng</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Yan</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Scale-aware fast R-CNN for pedestrian detection</article-title>. <source>IEEE Trans. Multimedia.</source> <volume>20</volume>, <fpage>985</fpage>&#x2013;<lpage>996</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/TMM.2017.2759508</pub-id>
</citation></ref>
<ref id="B18">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Liu</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Anguelov</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Erhan</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Szegedy</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Reed</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Fu</surname> <given-names>C.-Y.</given-names>
</name>
<etal/>
</person-group>. (<year>2016</year>). &#x201c;<article-title>Ssd: Single shot multibox detector</article-title>,&#x201d; in <conf-name>Computer Vision&#x2013;ECCV 2016: 14th European Conference</conf-name>, <conf-loc>Amsterdam, The Netherlands</conf-loc>, <conf-date>October 11&#x2013;14, 2016</conf-date>, Vol. <volume>14</volume>. <fpage>21</fpage>&#x2013;<lpage>37</lpage> (<publisher-loc>Amsterdam, Netherlands</publisher-loc>: <publisher-name>Springer</publisher-name>), Proceedings, Part I.</citation></ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Lu</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>First report of Colletotrichum truncatum causing anthracnose of lucky bamboo in Zhanjiang, China</article-title>. <source>Plant Dis.</source> <volume>103</volume>, <fpage>2947</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1094/PDIS-05-19-1122-PDN</pub-id>
</citation></ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Niu</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Zhong</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>H.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A review on the attention mechanism of deep learning</article-title>. <source>Neurocomputing</source> <volume>452</volume>, <fpage>48</fpage>&#x2013;<lpage>62</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.neucom.2021.03.091</pub-id>
</citation></ref>
<ref id="B21">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Pankaja</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Thippeswamy</surname> <given-names>G.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>Survey on leaf recognization and classification</article-title>,&#x201d; in <conf-name>2017 International Conference on Innovative Mechanisms for Industry Applications (ICIMIA)</conf-name>. <fpage>442</fpage>&#x2013;<lpage>450</lpage> (<publisher-loc>Bangalore, India</publisher-loc>: <publisher-name>IEEE</publisher-name>).</citation></ref>
<ref id="B22">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Redmon</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>You only look once: Unified, real-time object detection</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE conference on computer vision and pattern recognition</conf-name>. (<publisher-loc>Nevada, USA</publisher-loc>).</citation></ref>
<ref id="B23">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Redmon</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Farhadi</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>YOLO9000: better, faster, stronger</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE conference on computer vision and pattern recognition</conf-name>. <fpage>7263</fpage>&#x2013;<lpage>7271</lpage>. (<publisher-loc>Hawaii, USA</publisher-loc>).</citation></ref>
<ref id="B24">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Rezatofighi</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Tsoi</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Gwak</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Sadeghian</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Reid</surname> <given-names>I.</given-names>
</name>
<name>
<surname>Savarese</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>Generalized intersection over union: A metric and a loss for bounding box regression</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</conf-name>. <fpage>658</fpage>&#x2013;<lpage>666</lpage>. (<publisher-loc>California, USA</publisher-loc>).</citation></ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rout</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Mohapatra</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Jain</surname> <given-names>S. M.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>Tissue culture of ornamental pot plant: A critical review on present scenario and future prospects</article-title>. <source>Biotechnol. Adv.</source> <volume>24</volume>, <fpage>531</fpage>&#x2013;<lpage>560</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.biotechadv.2006.05.001</pub-id>, PMID: <pub-id pub-id-type="pmid">16814509</pub-id></citation></ref>
<ref id="B26">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Salvador</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Gir&#xf3;-i-Nieto</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Marqu&#xe9;s</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Satoh</surname> <given-names>S.i.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Faster r-cnn features for instance search</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE conference on computer vision and pattern recognition workshops</conf-name>. (<publisher-loc>Nevada, USA</publisher-loc>). <fpage>9</fpage>&#x2013;<lpage>16</lpage>.</citation></ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sawarkar</surname> <given-names>A. D.</given-names>
</name>
<name>
<surname>Shrimankar</surname> <given-names>D. D.</given-names>
</name>
<name>
<surname>Ali</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Agrahari</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Singh</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Bamboo plant classification using deep transfer learning with a majority multiclass voting algorithm</article-title>. <source>Appl. Sci.</source> <volume>14</volume>, <fpage>1023</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/app14031023</pub-id>
</citation></ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sharma</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Merritt</surname> <given-names>J. L.</given-names>
</name>
<name>
<surname>Palmateer</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Goss</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Smith</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Schubert</surname> <given-names>T.</given-names>
</name>
<etal/>
</person-group>. (<year>2014</year>). <article-title>Isolation, characterization, and management of Colletotrichum spp. causing anthracnose on lucky bamboo (Dracaena sanderiana)</article-title>. <source>HortScience</source> <volume>49</volume>, <fpage>453</fpage>&#x2013;<lpage>459</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.21273/HORTSCI.49.4.453</pub-id>
</citation></ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shorten</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Khoshgoftaar</surname> <given-names>T. M.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>A survey on image data augmentation for deep learning</article-title>. <source>J. Big. Data</source> <volume>6</volume>, <fpage>1</fpage>&#x2013;<lpage>48</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s40537-019-0197-0</pub-id>
</citation></ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tian</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Chu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Wei</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Shen</surname> <given-names>C.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Fully convolutional one-stage 3d object detection on lidar range images</article-title>. <source>Adv. Neural Inf. Process. Syst.</source> <volume>35</volume>, <fpage>34899</fpage>&#x2013;<lpage>34911</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2205.13764</pub-id>
</citation></ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tong</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Yu</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Wise-IoU: bounding box regression loss with dynamic focusing mechanism</article-title>. <source>arXiv. preprint. arXiv:.10051</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.48550/arXiv.2301.10051</pub-id>
</citation></ref>
<ref id="B32">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>van Dam</surname> <given-names>J. E.</given-names>
</name>
<name>
<surname>Elbersen</surname> <given-names>H. W.</given-names>
</name>
<name>
<surname>Monta&#xf1;o</surname> <given-names>C. M. D.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Bamboo production for industrial utilization</article-title>,&#x201d; in <source>Perennial grasses for bioenergy bioproducts</source>, <fpage>175</fpage>&#x2013;<lpage>216</lpage>. (<publisher-loc>Amsterdam, Netherlands</publisher-loc>: <publisher-name>Academic Press, Elsevier</publisher-name>)</citation></ref>
<ref id="B33">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>C.-Y.</given-names>
</name>
<name>
<surname>Bochkovskiy</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Liao</surname> <given-names>H.-Y. M.</given-names>
</name>
</person-group> (<year>2023</year>). &#x201c;<article-title>YOLOv7: Trainable bag-of-freebies sets new state-of-the-art for real-time object detectors</article-title>,&#x201d; in <conf-name>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</conf-name>. <fpage>7464</fpage>&#x2013;<lpage>7475</lpage>. (<publisher-loc>Vancouver, Canada</publisher-loc>).</citation></ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname> <given-names>Z.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Yue</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Yi</surname> <given-names>W.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>H.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>A phenomic approach of bamboo species identification using deep learning</article-title>. <source>Appl. Sci</source>. <volume>12</volume> (<issue>3</issue>), <fpage>1023</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.21203/rs.3.rs-1644335/v1</pub-id>
</citation></ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Watanabe</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Sumi</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Ise</surname> <given-names>T.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Identifying the vegetation type in Google Earth images using a convolutional neural network: a case study for Japanese bamboo forests</article-title>. <source>BMC Ecol.</source> <volume>20</volume>, <fpage>1</fpage>&#x2013;<lpage>14</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s12898-020-00331-5</pub-id>, PMID: <pub-id pub-id-type="pmid">33246473</pub-id></citation></ref>
</ref-list>
</back>
</article>