<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Bioeng. Biotechnol.</journal-id>
<journal-title>Frontiers in Bioengineering and Biotechnology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Bioeng. Biotechnol.</abbrev-journal-title>
<issn pub-type="epub">2296-4185</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">861286</article-id>
<article-id pub-id-type="doi">10.3389/fbioe.2022.861286</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Bioengineering and Biotechnology</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Real-Time Target Detection Method Based on Lightweight Convolutional Neural Network</article-title>
<alt-title alt-title-type="left-running-head">Yun et al.</alt-title>
<alt-title alt-title-type="right-running-head">Target Detection Based on Convolutional Neural Network</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Yun</surname>
<given-names>Juntong</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1513780/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Jiang</surname>
<given-names>Du</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1488063/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Liu</surname>
<given-names>Ying</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1513782/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Sun</surname>
<given-names>Ying</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1493953/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Tao</surname>
<given-names>Bo</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1488544/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Kong</surname>
<given-names>Jianyi</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Tian</surname>
<given-names>Jinrong</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1907970/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Tong</surname>
<given-names>Xiliang</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1431911/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Xu</surname>
<given-names>Manman</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1456167/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Fang</surname>
<given-names>Zifan</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1563760/overview"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Key Laboratory of Metallurgical Equipment and Control Technology of Ministry of Education</institution>, <institution>Wuhan University of Science and Technology</institution>, <addr-line>Wuhan</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Research Center for Biomimetic Robot and Intelligent Measurement and Control</institution>, <institution>Wuhan University of Science and Technology</institution>, <addr-line>Wuhan</addr-line>, <country>China</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Hubei Key Laboratory of Mechanical Transmission and Manufacturing Engineering</institution>, <institution>Wuhan University of Science and Technology</institution>, <addr-line>Wuhan</addr-line>, <country>China</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Precision Manufacturing Research Institute</institution>, <institution>Wuhan University of Science and Technology</institution>, <addr-line>Wuhan</addr-line>, <country>China</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Hubei Key Laboratory of Hydroelectric Machinery Design &#x26; Maintenance</institution>, <institution>China Three Gorges University</institution>, <addr-line>Yichang</addr-line>, <country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1254880/overview">Tinggui Chen</ext-link>, Zhejiang Gongshang University, China</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1657116/overview">Ashutosh Satapathy</ext-link>, Velagapudi Ramakrishna Siddhartha Engineering College, India</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1700705/overview">Maragatham G</ext-link>, SRM Institute of Science and Technology, India</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1488125/overview">Dongxu Gao</ext-link>, University of Portsmouth, United Kingdom</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1834940/overview">Teddy Surya Gunawan</ext-link>, International Islamic University Malaysia, Malaysia</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/388853/overview">Yinfeng Fang</ext-link>, Hangzhou Dianzi University, China</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Du Jiang, <email>jiangdu@wust.edu.cn</email>; Ying Liu, <email>liuying3025@wust.edu.cn</email>; Ying Sun, <email>sunying65@wust.edu.cn</email>; Zifan Fang, <email>fzf@ctgu.edu.cn</email>
</corresp>
<fn fn-type="other">
<p>This article was submitted to Bionics and Biomimetics, a section of the journal Frontiers in Bioengineering and Biotechnology</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>16</day>
<month>08</month>
<year>2022</year>
</pub-date>
<pub-date pub-type="collection">
<year>2022</year>
</pub-date>
<volume>10</volume>
<elocation-id>861286</elocation-id>
<history>
<date date-type="received">
<day>24</day>
<month>01</month>
<year>2022</year>
</date>
<date date-type="accepted">
<day>13</day>
<month>06</month>
<year>2022</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2022 Yun, Jiang, Liu, Sun, Tao, Kong, Tian, Tong, Xu and Fang.</copyright-statement>
<copyright-year>2022</copyright-year>
<copyright-holder>Yun, Jiang, Liu, Sun, Tao, Kong, Tian, Tong, Xu and Fang</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>The continuous development of deep learning improves target detection technology day by day. The current research focuses on improving the accuracy of target detection technology, resulting in the target detection model being too large. The number of parameters and detection speed of the target detection model are very important for the practical application of target detection technology in embedded systems. This article proposed a real-time target detection method based on a lightweight convolutional neural network to reduce the number of model parameters and improve the detection speed. In this article, the depthwise separable residual module is constructed by combining depthwise separable convolution and non&#x2013;bottleneck-free residual module, and the depthwise separable residual module and depthwise separable convolution structure are used to replace the VGG backbone network in the SSD network for feature extraction of the target detection model to reduce parameter quantity and improve detection speed. At the same time, the convolution kernels of 1 &#xd7; 3 and 3 &#xd7; 1 are used to replace the standard convolution of 3 &#xd7; 3 by adding the convolution kernels of 1 &#xd7; 3 and 3 &#xd7; 1, respectively, to obtain multiple detection feature graphs corresponding to SSD, and the real-time target detection model based on a lightweight convolutional neural network is established by integrating the information of multiple detection feature graphs. This article used the self-built target detection dataset in complex scenes for comparative experiments; the experimental results verify the effectiveness and superiority of the proposed method. The model is tested on video to verify the real-time performance of the model, and the model is deployed on the Android platform to verify the scalability of the model.</p>
</abstract>
<kwd-group>
<kwd>Deep learning</kwd>
<kwd>target detection</kwd>
<kwd>MobileNets-SSD</kwd>
<kwd>depthwise separable convolution</kwd>
<kwd>residual module</kwd>
</kwd-group>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>With the appearance and progress of powerful hardware devices such as image processors, deep learning has achieved rapid development. In recent years, deep convolutional neural networks have been widely applied to solve various tasks of computer vision. Traditional visual tasks include image classification, location, detection, and segmentation (<xref ref-type="bibr" rid="B50">Evan, et al., 2017</xref>; <xref ref-type="bibr" rid="B24">Jiang et al., 2019a</xref>; <xref ref-type="bibr" rid="B14">He, et al., 2019</xref>). In traditional visual tasks, feature extraction, a complicated task, has been completely replaced by convolutional neural networks (<xref ref-type="bibr" rid="B54">Sun, et al., 2020</xref>; <xref ref-type="bibr" rid="B60">Tian, et al., 2020</xref>; <xref ref-type="bibr" rid="B33">Liu, et al., 2021a</xref>; <xref ref-type="bibr" rid="B32">Liao, et al., 2021</xref>). On this basis, deep learning technology can improve the visual tasks of most complex scenes (<xref ref-type="bibr" rid="B28">Li, et al., 2019a</xref>; <xref ref-type="bibr" rid="B22">Jiang, et al., 2019b</xref>; <xref ref-type="bibr" rid="B17">Huang, et al., 2022</xref>). For example, automatic driving, face monitoring, pedestrian tracking, and so on are all tasks in very complex scenes, but the current research mostly focuses on how to improve the accuracy of target detection technology, which leads to the excessively large target detection model to a certain extent (<xref ref-type="bibr" rid="B6">Chen, et al., 2021a</xref>; <xref ref-type="bibr" rid="B3">Bai, et al., 2021</xref>; <xref ref-type="bibr" rid="B10">Duan, et al., 2021</xref>).</p>
<p>Target detection methods based on deep learning developed rapidly after 2012, which can be roughly divided into two categories: one is a two-stage model, which divides target detection into two stages: candidate box selection and target classification; the other is a one-stage model, which treats classification and localization as regression tasks. The two-stage target detection model first determines whether the target exists in the candidate region, and then determines the category with the classifier. However, most of the current research focuses on how to improve the accuracy of target detection technology, which leads to the excessively large target detection model to a certain extent. It is still challenging to synchronously realize high detection accuracy and real-time performance of objects in complex scenes.</p>
<p>This article proposes a real-time target detection method based on a lightweight convolutional neural network to reduce the parameters of the target detection model and improve the detection speed. First, Kinect is used to establish the target detection dataset in complex scenes, and the existing lightweight network is comprehensively studied. Then, combined with the depthwise convolution and bottleneck-free residual module, the depthwise residual module is proposed, and the MobileNet-SSD network is further improved by using the deep separable residual module, deep separable convolution, and convolution substitution structure. A real-time target detection model based on a lightweight convolutional neural network is established. The effectiveness of the proposed method is verified by comparing the established dataset with the existing lightweight target detection algorithm. Finally, the real-time detection model is tested on video, and the model is deployed to the mobile terminal to verify the scalability of the model.</p>
<p>The key contributions of this work are:<list list-type="simple">
<list-item>
<p>1) Combining depth-separable convolution and bottle-free residual module, the depth-separable residual module is proposed.</p>
</list-item>
<list-item>
<p>2) The MobileNet-SSD network is further improved by using the depthwise separable residual module, depthwise separable convolution, and convolutional substitution structure, and a real-time target detection method based on a lightweight convolutional neural network is proposed.</p>
</list-item>
<list-item>
<p>3) Target detection datasets are established in complex scenarios</p>
</list-item>
<list-item>
<p>4) Multiple groups of comparative experiments are conducted, and the proposed method is used to detect the video to verify the real-time performance of the model.</p>
</list-item>
</list>
</p>
<p>The rest of this article is organized as follows: <xref ref-type="sec" rid="s2">Section 2</xref> discusses the related work of target detection, followed by a target detection method based on improved MobileNet-SSD in <xref ref-type="sec" rid="s3">Section 3</xref>. A comparative experiment is carried out using self-built datasets in <xref ref-type="sec" rid="s4">Section 4</xref>. <xref ref-type="sec" rid="s5">Section 5</xref> concludes the paper with a summary and future research directions.</p>
</sec>
<sec id="s2">
<title>2 Related Work</title>
<p>At present, the mobile intelligent terminal has gradually become a necessity in people&#x2019;s life (<xref ref-type="bibr" rid="B25">Li, et al., 2019b</xref>; <xref ref-type="bibr" rid="B16">Hu, et al., 2019</xref>; <xref ref-type="bibr" rid="B67">Yu, et al., 2019</xref>; <xref ref-type="bibr" rid="B9">Cheng et al., 2021</xref>; <xref ref-type="bibr" rid="B20">Jiang, et al., 2021c</xref>; <xref ref-type="bibr" rid="B18">Huang, et al., 2021</xref>); while the mobile intelligent device for embedded devices is limited by the storage and computing power, the development of technology, such as unmanned drones, also need terminal real-time feedback image- and video-processing results; thus, the target detection model size and the complexity of calculation are difficult requirements (<xref ref-type="bibr" rid="B41">Luo, et al., 2020</xref>; <xref ref-type="bibr" rid="B36">Liu, et al., 2021b</xref>; <xref ref-type="bibr" rid="B55">Sun, et al., 2021</xref>; <xref ref-type="bibr" rid="B35">Liu, et al., 2022</xref>).</p>
<p>The task of target detection is to classify objects in the image and further determine their position in the image. For the recognition task, the network needs to extract deeper semantic features, that is, the essence of the target features, so as to distinguish between the target objects and improve the accuracy of recognition. For positioning tasks, location information needs to be saved as much as possible to bring the detection frame closer to the actual position of the target object in the image.</p>
<p>The traditional target detection process is as follows: first, multiple image regions with possible target objects are selected by sliding windows of different sizes; then, feature extraction methods such as SIFT (scale-invariant feature transform) and HOG (histogram of oriented gradient) are used to transform the information contained in the region into feature vectors and then classify them, commonly using the support vector machine (SVM) classifier. The DPM (deformable parts model) was proposed in 2010, which decomposes the target object into various parts for training and merges the prediction results of all parts during prediction to complete the detection of the target object. However, since the traditional target algorithm extracts the candidate region information and manually designs the features, the application range has great limitations. For example, the Haar feature is suitable for face detection, and the detector trained by this feature cannot detect other types of targets. In addition, the traditional target detection algorithm generates multiple candidate regions through traversal, which takes a lot of time. In addition, the traditional target detection algorithm classification training detector may produce the problem of feature vector &#x201c;dimension disaster.&#x201d;</p>
<p>Ross et al. proposed an R-CNN object detection model based on convolutional neural networks (CNNs), which first used depth to detect objects. However, the scaling of candidate regions has certain limitations in detection accuracy, and the training of this algorithm is complicated. In 2015, He et al. proposed the SPP-NET model to transform feature information of candidate regions of arbitrary size into feature vectors of fixed length. In the same year, Girshick proposed the fast R-CNN algorithm, which was based on ROI pooling (region of interest pooling), fixed the feature length of candidate regions, and used the multi-task loss function for training, which improved the training and detection efficiency of the target detection algorithm. In order to achieve real-time detection, researchers use the integrated convolutional neural network to complete target detection and improve the detection efficiency of the algorithm. Regression-based algorithms of YOLO and SSD (single-shot multibox detector) have continually appeared. However, both SSD and YOLO only use the characteristic information of a single scale for prediction, and the detection accuracy of multi-scale targets and small objects is low.</p>
<p>Due to the diversity of application scenarios of target detection technology <xref ref-type="bibr" rid="B52">(Sun, et al., 2022a</xref>; <xref ref-type="bibr" rid="B62">Weng, et al., 2021</xref>; <xref ref-type="bibr" rid="B69">Yun, et al., 2022</xref>; <xref ref-type="bibr" rid="B72">Zhao, et al., 2021</xref>), target detection algorithm should realize the lightweight of the model, solve the efficiency problem of the model, and successfully deploy or apply to mobile devices, industrial computers and other embedded platforms (<xref ref-type="bibr" rid="B64">Xiao, et al., 2021</xref>; <xref ref-type="bibr" rid="B66">Yang, et al., 2021</xref>; <xref ref-type="bibr" rid="B56">Sun, et al., 2022b</xref>; <xref ref-type="bibr" rid="B38">Liu et al., 2021c</xref>). Therefore, the lightweight target detection model has become another hot issue (<xref ref-type="bibr" rid="B42">Ma, et al., 2020</xref>; <xref ref-type="bibr" rid="B39">Liu, et al., 2021d</xref>). <xref ref-type="bibr" rid="B13">He et al. (2015)</xref> used the lightweight deep separable residual network as the basic network of fast R-CNN to reduce the parameters of the network model, fused the multi-layer convolution features in the basic network after local response normalization, enhanced the completeness of target feature information, and trained the network model in combination with Softmax loss function and central loss function so that the network model could learn other different target characteristics. <xref ref-type="bibr" rid="B47">Ren and Bao (2020)</xref> reduced the amount of network computation by using MobileNet as the basic network and replacing the standard convolution in the SSD detection layer with the inverse residual convolution. <xref ref-type="bibr" rid="B50">Evan et al. (2017)</xref> reduced darknet53, the backbone network of YOLOv3, and added an improved dense connection network and spatial pyramid pooling on the backbone network, which greatly improved the speed at the expense of accuracy. <xref ref-type="bibr" rid="B73">Zhao et al. (2020)</xref> integrated a 5 &#xd7; 5 depthwise separable convolution kernel on the basis of the MobileNetV2-SSD Lite model to further improve the recognition accuracy of the algorithm for small target objects, and the experimental results show that LMS-DN only needs fewer parameters and calculation costs to obtain higher identification accuracy and stronger anti-interference than other popular object detection models. <xref ref-type="bibr" rid="B70">Zhang et al. (2021)</xref> proposed a lightweight target detection network MN-YOLO (MobileNet-YOLOv4-tiny) suitable for embedded platforms using depthwise separable convolution instead of standard convolution to reduce the number of model parameters and calculations; at the same time, the visible light target detection model is used as the pretraining model of the infrared target detection model and the infrared target dataset collected on the spot is fine-tuned to obtain the infrared target detection model. Currently, miniaturized versions of YOLO and SSD algorithms are commonly used on embedded platforms (<xref ref-type="bibr" rid="B1">Alex, et al., 2017</xref>; <xref ref-type="bibr" rid="B4">Cao, et al., 2018</xref>; <xref ref-type="bibr" rid="B8">Cheng, et al., 2020</xref>; <xref ref-type="bibr" rid="B7">Chen, et al., 2021c</xref>; <xref ref-type="bibr" rid="B11">Hao, et al., 2021</xref>). The research of the MobileNet-SSD network framework to realize network model compression and multi-scale target detection is increasing gradually. Based on the Mobilenet-SSD framework, <xref ref-type="bibr" rid="B27">Li et al. (2019c)</xref> used the time characteristics of video to effectively improve the confidence level of detection and enhance the stability of detection, which provides a certain reference value for unmanned target detection. Although these algorithms have low computational load and fast detection speed, their detection accuracy is generally low, making it difficult to achieve a balance between computational load and accuracy (<xref ref-type="bibr" rid="B23">Jiang, et al., 2019d</xref>; <xref ref-type="bibr" rid="B21">Jiang et al., 2019e</xref>; <xref ref-type="bibr" rid="B19">Huang, et al., 2019</xref>; <xref ref-type="bibr" rid="B26">Li, et al., 2020</xref>).</p>
<p>To sum up, there are many algorithms for target detection at present, but the problems of target detection accuracy, model size, and detection speed still need to be solved in the application scenarios of service robots and other mobile devices (<xref ref-type="bibr" rid="B49">Sandler, et al., 2018</xref>; <xref ref-type="bibr" rid="B46">Qiu, et al., 2019</xref>; <xref ref-type="bibr" rid="B43">Meng, et al., 2020</xref>; <xref ref-type="bibr" rid="B68">Yu et al., 2020</xref>; <xref ref-type="bibr" rid="B30">Li, et al., 2021</xref>; <xref ref-type="bibr" rid="B58">Tao et al., 2022a</xref>). Therefore, a real-time target detection method based on a lightweight convolutional neural network is proposed in this article to reduce the number of target detection model parameters and improve the detection speed.</p>
</sec>
<sec id="s3">
<title>3 Improved MobileNet-SSD Network</title>
<sec id="s3-1">
<title>3.1 SSD</title>
<p>SSD is a one-stage target detection algorithm (<xref ref-type="bibr" rid="B57">Tan, et al., 2020</xref>; <xref ref-type="bibr" rid="B63">Wu, et al., 2022</xref>), which directly generates the category probability and position coordinate value of objects. After a single detection, the final detection result can be directly obtained, so it has a faster detection speed. The network detection framework is shown in <xref ref-type="fig" rid="F1">Figure 1</xref>. Traditional SSD uses VGG16 as the feature extraction network. The full connection layer of VGG16 is removed and the convolution layer is added to obtain more multi-layer feature maps for detection. At the same time, SSD makes full use of multi-level feature maps in the classification regression network, and the corresponding classification layer of all level feature maps shares weights with the location regression layer.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>SSD network structure.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g001.tif"/>
</fig>
<p>One of the cores of SSD is to detect objects of different sizes using feature maps of different levels, that is, to extract targets using feature maps output by each convolution layer. The scale of the anchor frame corresponding to the bottom-level feature graph to the high-rise feature graph is linearly divided from small to large. Steps for generating anchor frame are as follows:<list list-type="simple">
<list-item>
<p>1) A set of concentric anchor frames is generated centering on the midpoint of each point on the feature graph.</p>
</list-item>
<list-item>
<p>2) <inline-formula id="inf1">
<mml:math id="m1">
<mml:mi>m</mml:mi>
</mml:math>
</inline-formula> feature maps of different levels are used to extract targets. The scales of the bottom feature map corresponding to the anchor frame are <inline-formula id="inf2">
<mml:math id="m2">
<mml:mrow>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>min</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>, and the scales of the top are <inline-formula id="inf3">
<mml:math id="m3">
<mml:mrow>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>max</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>. That of the other layers are:</p>
</list-item>
</list>
<disp-formula id="e1">
<mml:math id="m4">
<mml:mrow>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>min</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>min</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>max</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mrow>
<mml:mo>[</mml:mo>
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>,</mml:mo>
<mml:mi>m</mml:mi>
</mml:mrow>
<mml:mo>]</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>
<list list-type="simple">
<list-item>
<p>3) Different ratios [1, 2, 3, 1/2, and 1/3] were used to calculate the width and height of the anchor frame using <xref ref-type="disp-formula" rid="e2">Eqs 2</xref>, <xref ref-type="disp-formula" rid="e3">3</xref>:</p>
</list-item>
</list>
<disp-formula id="e2">
<mml:math id="m5">
<mml:mrow>
<mml:msubsup>
<mml:mi>w</mml:mi>
<mml:mi>k</mml:mi>
<mml:mi>a</mml:mi>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
<mml:msqrt>
<mml:mrow>
<mml:msub>
<mml:mi>a</mml:mi>
<mml:mi>r</mml:mi>
</mml:msub>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:math>
<label>(2)</label>
</disp-formula>
<disp-formula id="e3">
<mml:math id="m6">
<mml:mrow>
<mml:msubsup>
<mml:mi>h</mml:mi>
<mml:mi>k</mml:mi>
<mml:mi>a</mml:mi>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
<mml:mo>/</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:msub>
<mml:mi>a</mml:mi>
<mml:mi>r</mml:mi>
</mml:msub>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:math>
<label>(3)</label>
</disp-formula>
<list list-type="simple">
<list-item>
<p>4) In the case of ratio &#x3d; 0, the specified scale is as follows:</p>
</list-item>
</list>
<disp-formula id="e4">
<mml:math id="m7">
<mml:mrow>
<mml:msubsup>
<mml:mi>s</mml:mi>
<mml:mi>k</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:math>
<label>(4)</label>
</disp-formula>
</p>
</sec>
<sec id="s3-2">
<title>3.2 MobileNet-SSD</title>
<p>The network detection framework of MobileNet-SSD is shown in <xref ref-type="fig" rid="F2">Figure 2</xref> (<xref ref-type="bibr" rid="B2">Algarni, 2021</xref>). The front-end network of MobileNet VGG16 is replaced by MobileNet, and the global average pooling layer, full connection layer, and Sofamax layer of MobileNet network are removed, followed by the back-end detection network of SSD. A MobileNet-SSD network was formed. Because the front-end network of the MobileNet-SSD network was deeper than that of SSD, the depth of the whole model was larger than that of the SSD network. From the perspective of the SSD back-end detection network, both MobileNet-SSD and SSD networks were detected by extracting features from the feature map of six scales. Because the MobileNet-SSD network adopted depthwise separable convolution, the resolution of the feature map of the back-end detection network was only half of that of the SSD network. Therefore, the network had less computation and computational complexity.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>MobileNet-SSD network structure.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g002.tif"/>
</fig>
</sec>
<sec id="s3-3">
<title>3.3 Improved MobileNet-SSD</title>
<p>The core of MobileNet is to consider image regions and channels separately and use depthwise convolution to replace standard convolution. The process of standard convolution is divided into depthwise convolution and pointwise convolution, that is, each channel is first convolved, then the information between channels is fused by 1 &#xd7; 1 convolution, the number of channels in the feature graph is changed, and the same effect as standard convolution is achieved (<xref ref-type="bibr" rid="B31">Liao, et al., 2020</xref>; <xref ref-type="bibr" rid="B40">Liu, et al., 2021e</xref>).</p>
<p>Depthwise separable convolution decomposes a complete convolution operation into two steps, that is, depthwise convolution and pointwise convolution. Different from conventional convolution operations, a convolution kernel of depthwise convolution is responsible for a channel, and a channel is convolved by only one convolution kernel. In the aforementioned conventional convolution, each convolution kernel operates on each channel of the input image simultaneously. Similarly, for a 128 &#xd7; 128 pixel, three-channel color input image (128 &#xd7; 128 &#xd7; 3), depthwise convolution intially goes through the first convolution operation. Different from the aforementioned conventional convolution, depthwise convolution is completely carried out on a two-dimensional plane. The number of convolution kernels is the same as the number of channels in the upper layer, that is, channels and convolution kernels correspond one to one. The operation of pointwise convolution is similar to that of conventional convolution operation. The size of its convolution kernel is <inline-formula id="inf4">
<mml:math id="m8">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, and <inline-formula id="inf5">
<mml:math id="m9">
<mml:mi>M</mml:mi>
</mml:math>
</inline-formula> is the number of channels in the upper layer. The convolution operation here will combine the feature graph of the previous step in the direction of the channel to generate a new feature graph.</p>
<p>The structure of standard convolution and depth-separable convolution is shown in <xref ref-type="fig" rid="F3">Figure 3</xref> (Liu, et al., 2021; Li, et al., 2019; <xref ref-type="bibr" rid="B11">Hao, et al., 2021</xref>), where the input image dimension is <inline-formula id="inf6">
<mml:math id="m10">
<mml:mrow>
<mml:mi>H</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>W</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> and the output image dimension is <inline-formula id="inf7">
<mml:math id="m11">
<mml:mrow>
<mml:mi>H</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>W</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>. The standard convolution can be obtained through the convolution kernel of <inline-formula id="inf8">
<mml:math id="m12">
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, and the required number of parameters is <inline-formula id="inf9">
<mml:math id="m13">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, while the depth-separable convolution is adopted. First, each channel of the input image is convolved, that is, the convolution kernel is <inline-formula id="inf10">
<mml:math id="m14">
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, and the required number of parameters is <inline-formula id="inf11">
<mml:math id="m15">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>. <inline-formula id="inf12">
<mml:math id="m16">
<mml:mi>M</mml:mi>
</mml:math>
</inline-formula> convolution 1 &#xd7; 1 is used to check the features of each channel for fusion. The number of parameters in this step is, N &#x00D7; k &#x00D7; k then the ratio of the number of parameters between the depthwise separable convolution and the standard convolution is shown in <xref ref-type="disp-formula" rid="e6">Eq. 5</xref>. When <inline-formula id="inf13">
<mml:math id="m17">
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>3</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>, the number of parameters of the depthwise separable convolution relative to the standard convolution is reduced by at least 8 to 9 times.<disp-formula id="e6">
<mml:math id="m18">
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>N</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>M</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>k</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mi>M</mml:mi>
</mml:mfrac>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mrow>
<mml:msup>
<mml:mi>k</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(5)</label>
</disp-formula>
</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Standard convolution and depthwise separable convolution. <bold>(A)</bold> Standard convolution. <bold>(B)</bold> Depthwise separable convolution.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g003.tif"/>
</fig>
<p>Two hyperparameters are set in MobileNet (<xref ref-type="bibr" rid="B18">Huang, et al., 2021</xref>), namely, the width multiplier and the resolution multiplier. The width multiplier controls the number of channels in the feature graph; when the width multiplier is less than 1, the model becomes thinner; the resolution multiplier is used to control the size of the feature graph, and both can reduce the number of parameters of the convolution flexibly. On the basis of MobileNet, MobileNetv2 uses an inverted residual block (Liu, et al., 2021; <xref ref-type="bibr" rid="B54">Sun, et al., 2020</xref>). First, 1 &#xd7; 1 convolution is used to improve the dimension of features, and then 3 &#xd7; 3 depth-separable convolution is used to extract features. Then, 1 &#xd7; 1 convolution is used to reduce dimensions.</p>
<p>The depthwise separable convolution network in MobileNet can greatly reduce the number of parameters in the network model. Therefore, the standard convolution in the VGG16 structure in SSD is replaced by the depthwise separable convolutional neural network. However, compared with the standard convolution, the network layers of the depthwise separable convolution are deeper. As the number of network layers increases, network performance degrades, that is, the detection accuracy begins to decline after reaching saturation. Therefore, in order to effectively solve the problem of network performance degradation, this article improved the MobileNet-SSD feature extraction network by combining the residual connection mode of the ResNet model and depthwise separable convolution.</p>
<p>If the input is set to <inline-formula id="inf14">
<mml:math id="m19">
<mml:mi>X</mml:mi>
</mml:math>
</inline-formula> and a parametrized network layer is set to <inline-formula id="inf15">
<mml:math id="m20">
<mml:mi>H</mml:mi>
</mml:math>
</inline-formula>, the output of this layer with <inline-formula id="inf16">
<mml:math id="m21">
<mml:mi>X</mml:mi>
</mml:math>
</inline-formula> as the input will be <inline-formula id="inf17">
<mml:math id="m22">
<mml:mrow>
<mml:mi>H</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mi>X</mml:mi>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>. General CNN networks, such as VGG, can directly learn the expression of parameter function <inline-formula id="inf18">
<mml:math id="m23">
<mml:mi>H</mml:mi>
</mml:math>
</inline-formula> through training, so as to directly learn <inline-formula id="inf19">
<mml:math id="m24">
<mml:mrow>
<mml:mi>X</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mo>&#x3e;</mml:mo>
<mml:mi>H</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mi>X</mml:mi>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>. Residual learning is committed to using multiple parametrized network layers to learn that the difference between input and output is <inline-formula id="inf20">
<mml:math id="m25">
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mi>H</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mi>X</mml:mi>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>X</mml:mi>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>X</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>. <inline-formula id="inf21">
<mml:math id="m26">
<mml:mi>X</mml:mi>
</mml:math>
</inline-formula> is the direct mapping, while <inline-formula id="inf22">
<mml:math id="m27">
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mi>H</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mi>X</mml:mi>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>X</mml:mi>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>X</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is the residual between input and output to be learned by the parameter network layer, and its principle is shown in <xref ref-type="fig" rid="F4">Figure 4</xref>.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Residual learning.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g004.tif"/>
</fig>
<p>The ResNet model has two types of residual modules, no-bottleneck residual module and bottleneck residual module, as shown in <xref ref-type="fig" rid="F5">Figure 5</xref>.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Two types of residual modules. <bold>(A)</bold> No-bottleneck residual module. <bold>(B)</bold> Bottleneck residual module.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g005.tif"/>
</fig>
<p>BN and Relu shown in <xref ref-type="fig" rid="F5">Figure 5</xref> are the normalization layer and activation function, respectively, which help to speed up the training and generalization of the network model. Compared with the no-bottleneck residual module, the bottleneck residual module uses 1 &#xd7; 1 convolution to reduce or expand the dimension of the feature graph, so that the 3 &#xd7; 3 convolution is no longer affected by the number of channels&#x2019; input, and accordingly, the output of this module will not affect the next module. The model layers are deep, and the bottleneck-free module is beneficial to improve the model detection accuracy, while the bottleneck residual module is beneficial to improve the model running speed.</p>
<p>Compared with the combination of depthwise separable convolution and bottleneck residual module, the combination of depthwise separable convolution and bottleneck residual module has more obvious advantages in reducing the number of model parameters. Therefore, the depthwise separable convolution is combined with the bottleneck-free residual module to improve the feature extraction function of the trunk network. The structure of the combined depthwise separable residual module is shown in <xref ref-type="fig" rid="F6">Figure 6</xref>. The network structure can effectively extract image feature information and greatly reduce the number of model parameters. Then, the module is combined with the depthwise separable structure to replace the VGG backbone network in the SSD network for feature extraction of the target detection model. Finally, for the network structure after Conv5_3 in SSD, the convolution sum of 1 &#xd7; 3 and 3 &#xd7; 1 convolution kernels are used to replace the standard convolution 3 &#xd7; 3, thus obtaining multiple detection feature graphs corresponding to SSD.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>The depthwise separable residual module structure.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g006.tif"/>
</fig>
<p>Both the bottleneck residual module and the non-bottleneck residual module can reduce the number of parameters and computation by introducing depth-separable convolution. <xref ref-type="table" rid="T1">Table 1</xref> compares the number of parameters of different types of residual modules when both input and output are 256 channels and 64 channels, respectively. In_Out_ C represents the number of input&#x2013;output channels, Bt represents the bottleneck residual module, Non-Bt represents the non-bottleneck residual module, DS-Bt represents the separable bottleneck residual module after the introduction of depthwise separable convolution, and DS-Non-Bt represents the separable bottleneck residual module after the introduction of depthwise separable convolution. When the input and output are 64 channels, the number of Bt parameters is 4.35K, the parameter of DS-Bt is 2.77K, the parameter of Non-Bt is 36.86K, and the parameter of DS-Non-Bt is 4.67K. The parameter number of DS-Bt is 63.7% of that of Bt, and that of DS-Non-Bt is 12.7% of that of Non-Bt. When the input and output channels are 256 channels, the parameter number of Bt is 69.63K, that of DS-Bt is 35.65K, that of Non-Bt is 589.82K, and that of DS-Non-Bt is 67.84K. The number of parameters of DS-Bt is 51.2% of that of Bt, and that of DS-Non-Bt is 11.5% of that of Non-Bt. It can be seen from these data that the depth-separable convolution introduced by the bottleneck residual module has a higher benefit in reducing the number of parameters than the depth-separable convolution introduced by the bottleneck residual module. Moreover, the more channels there are, the more benefit can be obtained in reducing the number of parameters by introducing depthwise separable convolution.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Number of module parameters with different residuals.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Residual block</th>
<th align="left"/>
<th align="center">Bt (K)</th>
<th align="center">Non-Bt (K)</th>
<th align="center">DS-Bt (K)</th>
<th align="center">DS-non-Bt (K)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td rowspan="2" align="left">In_Out_C</td>
<td align="char" char=".">64</td>
<td align="char" char=".">4.35</td>
<td align="char" char=".">36.86</td>
<td align="char" char=".">2.77</td>
<td align="char" char=".">4.67</td>
</tr>
<tr>
<td align="char" char=".">256</td>
<td align="char" char=".">69.63</td>
<td align="char" char=".">589.82</td>
<td align="char" char=".">35.65</td>
<td align="char" char=".">67.84</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>The specific parameters of lightweight SSD network structure based on depthwise separable convolution are shown in <xref ref-type="table" rid="T2">Tables 2</xref> and <xref ref-type="table" rid="T3">3</xref>, where Conv is the standard convolution, DW is the depthwise separable convolution, DS-RES is the depthwise separable residual module, and Alter Conv is the alternative convolution of corresponding parameters. The improved SSD adopts the idea of multi-layer feature detection in SSD. Multiple DS-RES modules are used to extract features, and use the feature graph of 19 &#xd7; 19, 10 &#xd7; 10, 5 &#xd7; 5, 3 &#xd7; 3, 2 &#xd7; 2, and 1 &#xd7; 1 for detection.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>The structure of a real-time target detection algorithm based on a lightweight convolutional neural network.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Network layer</th>
<th align="center">Output size</th>
<th align="center">Convolution Kernel size</th>
<th align="center">Step</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Input</td>
<td align="center">300 &#xd7; 300 &#xd7; 3</td>
<td align="left"/>
<td align="left"/>
</tr>
<tr>
<td align="left">Conv1</td>
<td align="center">150 &#xd7; 150 &#xd7; 32</td>
<td align="center">3 &#xd7; 3,32</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">DW1</td>
<td align="center">150 &#xd7; 150 &#xd7; 64</td>
<td align="center">3 &#xd7; 3,64</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">DS-Res2</td>
<td align="center">150 &#xd7; 150 &#xd7; 64</td>
<td align="center">3 &#xd7; 3,64 3 &#xd7; 3,64</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">DW3</td>
<td align="center">75 &#xd7; 75 &#xd7; 128</td>
<td align="center">1 &#xd7; 1,128</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">DW4</td>
<td align="center">75 &#xd7; 75 &#xd7; 128</td>
<td align="center">1 &#xd7; 1,128</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">DS-Res5</td>
<td align="center">75 &#xd7; 75 &#xd7; 128</td>
<td align="center">3 &#xd7; 3,128 3 &#xd7; 3,128</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">DW6</td>
<td align="center">38 &#xd7; 38 &#xd7; 256</td>
<td align="center">1 &#xd7; 1,256</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">DW7</td>
<td align="center">38 &#xd7; 38 &#xd7; 256</td>
<td align="center">1 &#xd7; 1,256</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">DS-Res8</td>
<td align="center">38 &#xd7; 38 &#xd7; 256</td>
<td align="center">3 &#xd7; 3,256 3 &#xd7; 3,256</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">DW9</td>
<td align="center">19 &#xd7; 19 &#xd7; 512</td>
<td align="center">1 &#xd7; 1,512</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">DW (10&#x2013;14)</td>
<td align="center">19 &#xd7; 19 &#xd7; 512</td>
<td align="center">(1 &#xd7; 1,256)&#xd7;5</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">DS-Res15</td>
<td align="center">19 &#xd7; 19 &#xd7; 512</td>
<td align="center">3 &#xd7; 3,512 3 &#xd7; 3,512</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">DW16</td>
<td align="center">10 &#xd7; 10 &#xd7; 1024</td>
<td align="center">1 &#xd7; 1,1024</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">DW17</td>
<td align="center">10 &#xd7; 10 &#xd7; 1024</td>
<td align="center">1 &#xd7; 1,1024</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">Conv2</td>
<td align="center">10 &#xd7; 10 &#xd7; 256</td>
<td align="center">1 &#xd7; 1,256</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">Alter Conv1</td>
<td align="center">5 &#xd7; 5 &#xd7; 256</td>
<td align="center">3 &#xd7; 3,256</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">Conv3</td>
<td align="center">5 &#xd7; 5 &#xd7; 128</td>
<td align="center">1 &#xd7; 1,128</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">Alter Conv2</td>
<td align="center">3 &#xd7; 3 &#xd7; 256</td>
<td align="center">3 &#xd7; 3,256</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">Conv4</td>
<td align="center">3 &#xd7; 3 &#xd7; 128</td>
<td align="center">1 &#xd7; 1,128</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">Alter Conv3</td>
<td align="center">2 &#xd7; 2 &#xd7; 256</td>
<td align="center">3 &#xd7; 3,256</td>
<td align="center">2</td>
</tr>
<tr>
<td align="left">Conv5</td>
<td align="center">2 &#xd7; 2 &#xd7; 64</td>
<td align="center">1 &#xd7; 1,64</td>
<td align="center">1</td>
</tr>
<tr>
<td align="left">Alter Conv5</td>
<td align="center">1 &#xd7; 1 &#xd7; 128</td>
<td align="center">3 &#xd7; 3,128</td>
<td align="center">2</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Parameters related to the experimental environment.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Category name</th>
<th align="center">Parameter</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">operating system</td>
<td>Windows 10</td>
</tr>
<tr>
<td align="left">CPU</td>
<td>AMD Ryzen 7</td>
</tr>
<tr>
<td align="left">GPU</td>
<td>NVIDIA GeForce RTX 2070</td>
</tr>
<tr>
<td align="left">Cuda with Cudnn</td>
<td>10.0/7.6.5</td>
</tr>
<tr>
<td align="left">Python</td>
<td>3.6</td>
</tr>
<tr>
<td align="left">Tensorflow, Keras</td>
<td>1.13.2/2.1.5</td>
</tr>
<tr>
<td align="left">Opencv</td>
<td>4.5.1</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>The loss function is the weighted sum of position error and confidence error, as shown in <xref ref-type="disp-formula" rid="e7">Eq. 6</xref>.<disp-formula id="e7">
<mml:math id="m28">
<mml:mrow>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>c</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>g</mml:mi>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:mfrac>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>f</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>c</mml:mi>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>&#x3b1;</mml:mi>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mi>l</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>c</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>g</mml:mi>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(6)</label>
</disp-formula>where,<disp-formula id="e8">
<mml:math id="m29">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mi>l</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>c</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>g</mml:mi>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>x</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>c</mml:mi>
<mml:mi>y</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>w</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>h</mml:mi>
</mml:mrow>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:munder>
<mml:mrow>
<mml:msubsup>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mi>k</mml:mi>
</mml:msubsup>
<mml:mi>s</mml:mi>
<mml:mi>m</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>t</mml:mi>
<mml:msub>
<mml:mi>h</mml:mi>
<mml:mrow>
<mml:mi>L</mml:mi>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msubsup>
<mml:mi>l</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>m</mml:mi>
</mml:msubsup>
<mml:mo>&#x2212;</mml:mo>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mi>m</mml:mi>
</mml:msubsup>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:math>
<label>(7)</label>
</disp-formula>where,<disp-formula id="e9">
<mml:math id="m30">
<mml:mrow>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x2212;</mml:mo>
<mml:msubsup>
<mml:mi>d</mml:mi>
<mml:mi>i</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:msubsup>
<mml:mi>d</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>w</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mtext>&#x2002;</mml:mtext>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>y</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>y</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x2212;</mml:mo>
<mml:msubsup>
<mml:mi>d</mml:mi>
<mml:mi>i</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>y</mml:mi>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:msubsup>
<mml:mi>d</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>h</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(8)</label>
</disp-formula>
<disp-formula id="e10">
<mml:math id="m31">
<mml:mrow>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mi>w</mml:mi>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>log</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mi>w</mml:mi>
</mml:msubsup>
</mml:mrow>
<mml:mrow>
<mml:msubsup>
<mml:mi>d</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>w</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>,</mml:mo>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mtext>&#x2002;</mml:mtext>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mi>h</mml:mi>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>log</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mi>h</mml:mi>
</mml:msubsup>
</mml:mrow>
<mml:mrow>
<mml:msubsup>
<mml:mi>d</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>h</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(9)</label>
</disp-formula>where <inline-formula id="inf23">
<mml:math id="m32">
<mml:mi>N</mml:mi>
</mml:math>
</inline-formula> is the number of prior frames of positive samples; <inline-formula id="inf24">
<mml:math id="m33">
<mml:mi>&#x3b1;</mml:mi>
</mml:math>
</inline-formula> is the weight coefficient, set as 1; <inline-formula id="inf25">
<mml:math id="m34">
<mml:mrow>
<mml:msubsup>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mi>p</mml:mi>
</mml:msubsup>
<mml:mo>&#x2208;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mrow>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>, when <inline-formula id="inf26">
<mml:math id="m35">
<mml:mrow>
<mml:msubsup>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mi>p</mml:mi>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>, it means that the prior box <inline-formula id="inf27">
<mml:math id="m36">
<mml:mi>i</mml:mi>
</mml:math>
</inline-formula> matches the target <inline-formula id="inf28">
<mml:math id="m37">
<mml:mi>j</mml:mi>
</mml:math>
</inline-formula> and the target category is <inline-formula id="inf29">
<mml:math id="m38">
<mml:mi>p</mml:mi>
</mml:math>
</inline-formula>; <inline-formula id="inf30">
<mml:math id="m39">
<mml:mi>c</mml:mi>
</mml:math>
</inline-formula> is the predicted value of category confidence; <inline-formula id="inf31">
<mml:math id="m40">
<mml:mi>l</mml:mi>
</mml:math>
</inline-formula> is the position prediction value of prior frame; <inline-formula id="inf32">
<mml:math id="m41">
<mml:mi>g</mml:mi>
</mml:math>
</inline-formula> is the location parameter of the real target; and <inline-formula id="inf33">
<mml:math id="m42">
<mml:mrow>
<mml:msubsup>
<mml:mi>g</mml:mi>
<mml:mi>j</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> is the encoding of the real box.</p>
<p>The confidence error is a Softmax function:<disp-formula id="e11">
<mml:math id="m43">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>f</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>,</mml:mo>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mi>c</mml:mi>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mo>&#x2212;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
<mml:mrow>
<mml:msubsup>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mi>p</mml:mi>
</mml:msubsup>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>log</mml:mi>
<mml:mo>&#x2061;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#x2227;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>p</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:mstyle>
<mml:mo>&#x2212;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mi>N</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>g</mml:mi>
</mml:mrow>
</mml:munder>
<mml:mrow>
<mml:mi>log</mml:mi>
<mml:mo>&#x2061;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#x2227;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>o</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:mstyle>
<mml:mo>,</mml:mo>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mtext>&#x2002;</mml:mtext>
<mml:mi>w</mml:mi>
<mml:mi>h</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#x2002;</mml:mtext>
<mml:msubsup>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#x2227;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>p</mml:mi>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msubsup>
<mml:mi>c</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>p</mml:mi>
</mml:msubsup>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mi>p</mml:mi>
</mml:munder>
<mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msubsup>
<mml:mi>c</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>p</mml:mi>
</mml:msubsup>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>.</mml:mo>
</mml:math>
<label>(10)</label>
</disp-formula>
</p>
</sec>
</sec>
<sec id="s4">
<title>4 Experiment and Analysis</title>
<sec id="s4-1">
<title>4.1 Establishment of Target Detection Dataset in Complex Scenarios</title>
<p>The Kinect camera was used to collect studio scenes in the manner of a video stream, and common objects in daily life were selected as detection targets, including toys, chair, stool, cabinet, glasses case, and cup. In the process of image collection, 1,064 color images of studio indoor scenes with different backgrounds, different light intensity, and different angles were collected, and the deformation of the toy page, thermos cup, and glasses case with different poses was taken into account. The chair and stool shape similarity improved the robustness of the target detection model. The collected pictures were named in one-to-one correspondence with four Arabic digits, and part of the sample of the indoor scene image constructed from this is shown in <xref ref-type="fig" rid="F7">Figure 7</xref>.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Color images of different angles, backgrounds, and lighting.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g007.tif"/>
</fig>
<p>Although the established image database contained images in various scenarios, the samples still lacked diversity. Therefore, on the basis of the established image data set, in order to increase the noise anti-interference ability of the model, the image of the dataset random chose some image processing operations; to make the data richer, each category contained a sample generally reaching equilibrium level so that it could be used to enhance the training dataset of the network, get better model performance, and improve the generalizability. Therefore, under the condition that other conditions remain unchanged, random rotation transform, inversion transform, image translation transform, noise disturbance, random clipping transform, image color transform, random occlusion, and random superposition of the aforementioned operations were carried out on the collected images to expand the dataset to 4,256 pieces. Label-Img was used to annotate the image dataset by category and position, and the indoor scene dataset was created.</p>
</sec>
<sec id="s4-2">
<title>4.2 Experiment and Result Analysis</title>
<p>In this article, the improved MobileNet-SSD was trained by using the target detection dataset in complex scenarios. The parameter configuration of the experimental environment is shown in <xref ref-type="table" rid="T2">Tables 2</xref> and <xref ref-type="table" rid="T3">3</xref>. The Adam optimizer was used to adjust the learning rate during the training process. The training situations shown in <xref ref-type="fig" rid="F8">Figures 8A,B</xref> represent the loss of the training set and verification set in the training process, respectively.</p>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>Training of trial target detection model based on a lightweight convolutional neural network. <bold>(A)</bold> Training set loss. <bold>(B)</bold> Validation set loss.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g008.tif"/>
</fig>
<p>The comparative experiment is conducted on SSD, Tiny-Yolov3, Mobilenet-SSD, and the improved MobileNet-SSD on the complex scene dataset. The detection of each algorithm for each category is shown in <xref ref-type="fig" rid="F9">Figure 9</xref>.</p>
<fig id="F9" position="float">
<label>FIGURE 9</label>
<caption>
<p>Comparison of detection accuracy between SSD and lightweight target detection algorithms for various classes. <bold>(A)</bold> SSD. <bold>(B)</bold> Improved MobileNet-SSD. <bold>(C)</bold> MobileNet-SSD. <bold>(D)</bold> Tiny-YOLOv3.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g009.tif"/>
</fig>
<p>The comparison between the detection accuracy, speed, model parameters, and training time of SSD and several lightweight target detection algorithms is shown in <xref ref-type="table" rid="T4">Table 4</xref>. As can be seen from the table, compared with SSD, the detection accuracy of SSD improved by using the depthwise separable residual module was not reduced, but the number of model parameters was greatly reduced, which is conducive to model deployment, improves detection speed, and improves the real-time performance of the target detection algorithm. Compared with Mobilenet-SSD and Tiny-YOLOv3, SSD based on depthwise separable convolution had a smaller number of model parameters and a lower detection speed, but had a huge advantage in detection accuracy. When the confidence threshold is set to 0.5, the detection effect of SSD, lightweight SSD, Mobilenet-SSD, and Tiny-YOLOv3 on the same image is shown in <xref ref-type="fig" rid="F10">Figure 10</xref>.</p>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Performance comparison between SSD and lightweight target detection algorithms.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Evaluation standard algorithm</th>
<th align="center">mAP, %</th>
<th align="center">FPS</th>
<th align="center">MB</th>
<th align="center">Training time/min</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">SSD</td>
<td align="char" char=".">87.13</td>
<td align="char" char=".">26</td>
<td align="char" char=".">93.2</td>
<td align="char" char=".">37</td>
</tr>
<tr>
<td align="left">Improved MobileNet-SSD</td>
<td align="char" char=".">87.33</td>
<td align="char" char=".">47</td>
<td align="char" char=".">27.3</td>
<td align="char" char=".">12.6</td>
</tr>
<tr>
<td align="left">Tiny-YOLOv3</td>
<td align="char" char=".">66.57</td>
<td align="char" char=".">52</td>
<td align="char" char=".">33.2</td>
<td align="char" char=".">15.3</td>
</tr>
<tr>
<td align="left">MobileNet-SSD</td>
<td align="char" char=".">67.02</td>
<td align="char" char=".">62</td>
<td align="char" char=".">26.8</td>
<td align="char" char=".">11.4</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F10" position="float">
<label>FIGURE 10</label>
<caption>
<p>Comparison of detection effects between SSD and the lightweight target detection model. <bold>(A)</bold> SSD. <bold>(B)</bold> Improved MobileNet-SSD. <bold>(C)</bold> MobileNet-SSD. <bold>(D)</bold> Tiny-YOLOv3.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g010.tif"/>
</fig>
<p>The real-time detection model was tested on video, and its detection speed met the real-time requirement. <xref ref-type="fig" rid="F11">Figure 11</xref> shows the detection effect of the real-time target detection model on video.</p>
<fig id="F11" position="float">
<label>FIGURE 11</label>
<caption>
<p>Detection effect of real-time detection model on video.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g011.tif"/>
</fig>
<p>It has become a trend for the model to run on the mobile terminal. In order to verify the scalability of the model, the TensorFlow model generated by Android Studio was deployed to the Android mobile terminal, the project was compiled and run, the deployment of the real-time and high-precision target detection model on the mobile end was completed, and the real-time detection on the mobile end was realized. The experimental results are shown in <xref ref-type="fig" rid="F12">Figure 12</xref>.</p>
<fig id="F12" position="float">
<label>FIGURE 12</label>
<caption>
<p>Deployment of real-time detection model on the Android platform.</p>
</caption>
<graphic xlink:href="fbioe-10-861286-g012.tif"/>
</fig>
</sec>
</sec>
<sec id="s5">
<title>5 Conclusion</title>
<p>In order to solve the application problem of the target detection model in embedded devices and mobile terminals, this article focuses on the research of target detection algorithm lightweight. First, the MobileNet-SSD network was introduced and analyzed, and then improved by combining the depthwise separable convolution, no-bottleneck residual module, and the convolution substitution structure to reduce parameter quantity and improve detection speed. A comparative experiment was carried out on the self-built complex scene target detection dataset; the experimental results show that the MobileNet-SSD improved relative to the SSD model precision without loss and greatly reduced the number of parameters of the model, which is advantageous to the model in the mobile terminal, deployment of embedded devices, and improvement of the detection speed of the algorithm, namely, the real-time target detection. Compared with the existing lightweight target detection network, the real-time target detection model based on the lightweight convolutional neural network proposed in this article has similar parameters, but has great advantages in detection accuracy. Finally, the model was tested on video to verify the real-time performance of the model, and the model is deployed on the Android platform to verify the scalability of the model. There are still shortcomings in this study. In future research, the neural structure search method can be used to optimize the detection speed and accuracy of the model while limiting the number of neural network parameters, so as to achieve high accuracy and real-time performance of target detection technology on embedded devices.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s6">
<title>Data Availability Statement</title>
<p>The original contributions presented in the study are included in the article/Supplementary Material; further inquiries can be directed to the corresponding authors.</p>
</sec>
<sec id="s7">
<title>Author Contributions</title>
<p>JY and DJ provided research ideas and plans; YL and YS wrote programs and conducted experiments; BT, JK, and JT analyzed and explained the simulation results; XT improved the algorithm. MX and ZF co-authored the manuscript, JY and DJ were responsible for collecting data; and DJ, YS, YL, and ZF revised the manuscript for the corresponding author and approved the final submission.</p>
</sec>
<sec id="s8">
<title>Funding</title>
<p>This work is supported by grants from the National Natural Science Foundation of China (Grant Nos. 52075530,51575407, 51975324, 51505349, 61733011, and 41906177); the Grants of Hubei Provincial Department of Education (D20191105); the Grants of National Defense PreResearch Foundation of Wuhan University of Science and Technology (GF201705), Open Fund of the Key Laboratory for Metallurgical Equipment and Control of Ministry of Education in Wuhan University of Science and Technology (2018B07 and 2019B13), and Open Fund of Hubei Key Laboratory of Hydroelectric Machinery Design &#x26; Maintenance in China Three Gorges University(2020KJX02 and 2021KJX13).</p>
</sec>
<sec sec-type="COI-statement" id="s9">
<title>Conflict of Interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x2019;s Note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors, and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Alex</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Ilya</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Geoffrey</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>ImageNet Classification with Deep Convolutional Neural Networks</article-title>. <source>Commun. ACM</source> <volume>60</volume> (<issue>6</issue>), <fpage>84</fpage>&#x2013;<lpage>90</lpage>. <pub-id pub-id-type="doi">10.1145/3065386</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1145/3065386">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=ImageNet+Classification+with+Deep+Convolutional+Neural+Networks&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Algarni</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A Lightweight Cryptography (LWC) Framework to Secure Memory Heap in Internet of Things</article-title>. <source>Alexandria Eng. J.</source> <volume>60</volume> (<issue>1</issue>), <fpage>1489</fpage>&#x2013;<lpage>1497</lpage>. <pub-id pub-id-type="doi">10.1016/j.aej.2020.11.003</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.aej.2020.11.003">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=A+Lightweight+Cryptography+(LWC)+Framework+to+Secure+Memory+Heap+in+Internet+of+Things&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bai</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Impro-ved Single Shot Multibox Detector Target Detection Method Based on Deep Feature Fusion</article-title>. <source>Concurrency Comput. Pract. Exp</source> <volume>34</volume> (<issue>4</issue>), <fpage>e6614</fpage>. <pub-id pub-id-type="doi">10.1002/cpe.6614</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/cpe.6614">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Impro-ved+Single+Shot+Multibox+Detector+Target+Detection+Method+Based+on+Deep+Feature+Fusion&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B4">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Cao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2018</year>). <source>Low Altitude Armored Target Detection Based on Rotation Invariant Faster R-CNN</source>. <publisher-loc>Shanghai, China.</publisher-loc>: <publisher-name>Laser &#x26; Optoelectronics Progress</publisher-name>. <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Low+Altitude+Armored+Target+Detection+Based+on+Rotation+Invariant+Faster+R-CNN&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Cong</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2021b</year>). <article-title>Analysis of User Needs on Downloading Behavior of English Vocabulary APPs Based on Data Mining for Online Comments</article-title>. <source>Mathematics</source> <volume>9</volume> (<issue>12</issue>), <fpage>1341</fpage>. <pub-id pub-id-type="doi">10.3390/math9121341</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3390/math9121341">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Analysis+of+User+Needs+on+Downloading+Behavior+of+English+Vocabulary+APPs+Based+on+Data+Mining+for+Online+Comments&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Rong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Cong</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2021a</year>). <article-title>Combining Public Opinion Dissemination with Polarization Process Considering Individual Heterogeneity</article-title>. <source>Healthcare</source> <volume>9</volume> (<issue>2</issue>), <fpage>176</fpage>. <pub-id pub-id-type="doi">10.3390/healthcare9020176</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3390/healthcare9020176">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Combining+Public+Opinion+Dissemination+with+Polarization+Process+Considering+Individual+Heterogeneity&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Yin</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>l.</given-names>
</name>
<name>
<surname>Rong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Cong</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2021c</year>). <article-title>Monitoring and Recognizing Enterprise Public Opinion from High-Risk Users Based on User Portrait and Random Forest Algorithm</article-title>. <source>Axioms</source> <volume>10</volume> (<issue>2</issue>), <fpage>106</fpage>. <pub-id pub-id-type="doi">10.3390/axioms10020106</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3390/axioms10020106">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Monitoring+and+Recognizing+Enterprise+Public+Opinion+from+High-Risk+Users+Based+on+User+Portrait+and+Random+Forest+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Zeng</surname>
<given-names>F.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Visualization of Activated Muscle Area Based on sEMG</article-title>. <source>Ifs</source> <volume>38</volume> (<issue>3</issue>), <fpage>2623</fpage>&#x2013;<lpage>2634</lpage>. <pub-id pub-id-type="doi">10.3233/jifs-179549</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3233/jifs-179549">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Visualization+of+Activated+Muscle+Area+Based+on+sEMG&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Gesture Recognition Based on Surface Electromyography &#x2010;feature Image</article-title>. <source>Concurr. Comput. Pract. Exper</source> <volume>33</volume> (<issue>6</issue>), <fpage>e6051</fpage>. <pub-id pub-id-type="doi">10.1002/cpe.6051</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/cpe.6051">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Gesture+Recognition+Based+on+Surface+Electromyography+&#x2010;feature+Image&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Duan</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Gesture Recognition Based on Multi&#x2010;modal Feature Weight</article-title>. <source>Concurr. Comput. Pract. Exper</source> <volume>33</volume> (<issue>5</issue>), <fpage>e5991</fpage>. <pub-id pub-id-type="doi">10.1002/cpe.5991</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/cpe.5991">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Gesture+Recognition+Based+on+Multi&#x2010;modal+Feature+Weight&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hao</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Bai</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Intelligent Detection of Steel Defects Based on Improved Split Attention Networks</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>9</volume>, <fpage>810876</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2021.810876</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35096796/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2021.810876">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Intelligent+Detection+of+Steel+Defects+Based+on+Improved+Split+Attention+Networks&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hao</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Bai</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Towards the Steel Plate Defect Detection: Multidimensional Feature Information Extraction and Fusion</article-title>. <source>Concurr. Comput. Pract. Exper</source> <volume>33</volume> (<issue>21</issue>), <fpage>e6384</fpage>. <pub-id pub-id-type="doi">10.1002/CPE.6384</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/CPE.6384">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Towards+the+Steel+Plate+Defect+Detection:+Multidimensional+Feature+Information+Extraction+and+Fusion&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Ren</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Spatial Pyramid Pooling in Deep Convolutional Networks for Visual Recognition</article-title>. <source>IEEE Trans. Pattern Anal. Mach. Intell.</source> <volume>37</volume> (<issue>9</issue>), <fpage>1904</fpage>&#x2013;<lpage>1916</lpage>. <pub-id pub-id-type="doi">10.1109/tpami.2015.2389824</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/26353135/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/tpami.2015.2389824">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Spatial+Pyramid+Pooling+in+Deep+Convolutional+Networks+for+Visual+Recognition&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Liao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Gesture Recognition Based on an Improved Local Sparse Representation Classification Algorithm</article-title>. <source>Clust. Comput.</source> <volume>22</volume> (<issue>Suppl. 5</issue>), <fpage>10935</fpage>&#x2013;<lpage>10946</lpage>. <pub-id pub-id-type="doi">10.1007/s10586-017-1237-1</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s10586-017-1237-1">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Gesture+Recognition+Based+on+an+Improved+Local+Sparse+Representation+Classification+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Probability Analysis for Grasp Planning Facing the Field of Medical Robotics</article-title>. <source>Measurement</source> <volume>141</volume>, <fpage>227</fpage>&#x2013;<lpage>234</lpage>. <pub-id pub-id-type="doi">10.1016/j.measurement.2019.03.010</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.measurement.2019.03.010">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Probability+Analysis+for+Grasp+Planning+Facing+the+Field+of+Medical+Robotics&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Hao</surname>
<given-names>Z.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Multi-scale Feature Fusion Convolutional Neural Network for Indoor Small Target Detection</article-title>. <source>Front. Neurorobot.</source> <volume>16</volume>, <fpage>881021</fpage>. <pub-id pub-id-type="doi">10.3389/fnbot.2022.881021</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35663726/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fnbot.2022.881021">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Multi-scale+Feature+Fusion+Convolutional+Neural+Network+for+Indoor+Small+Target+Detection&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Fu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Hao</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Detection Algorithm of Safety Helmet Wearing Based on Deep Learning</article-title>. <source>Concurr. Comput. Pract. Exper</source> <volume>33</volume> (<issue>13</issue>), <fpage>e6234</fpage>. <pub-id pub-id-type="doi">10.1002/CPE.6234</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/CPE.6234">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Detection+Algorithm+of+Safety+Helmet+Wearing+Based+on+Deep+Learning&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Fu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Luo</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Improvement of Maximum Variance Weight Partitioning Particle Filter in Urban Computing and Intelligence</article-title>. <source>IEEE Access</source> <volume>7</volume>, <fpage>106527</fpage>&#x2013;<lpage>106535</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2019.2932144</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/ACCESS.2019.2932144">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Improvement+of+Maximum+Variance+Weight+Partitioning+Particle+Filter+in+Urban+Computing+and+Intelligence&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2021c</year>). <article-title>Manipulator Grabbing Position Detection with Information Fusion of Color Image and Depth Image Using Deep Learning</article-title>. <source>J. Ambient. Intell. Hum. Comput.</source> <volume>12</volume> (<issue>12</issue>), <fpage>10809</fpage>&#x2013;<lpage>10822</lpage>. <pub-id pub-id-type="doi">10.1007/s12652-020-02843-w</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s12652-020-02843-w">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Manipulator+Grabbing+Position+Detection+with+Information+Fusion+of+Color+Image+and+Depth+Image+Using+Deep+Learning&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2019e</year>). <article-title>Grip Strength Forecast and Rehabilitative Guidance Based on Adaptive Neural Fuzzy Inference System Using sEMG</article-title>. <source>Pers. Ubiquit Comput</source>. <pub-id pub-id-type="doi">10.1007/s00779-019-01268-3</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s00779-019-01268-3">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Grip+Strength+Forecast+and+Rehabilitative+Guidance+Based+on+Adaptive+Neural+Fuzzy+Inference+System+Using+sEMG&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2019b</year>). <article-title>Gesture Recognition Based on Skeletonization Algorithm and CNN with ASL Database</article-title>. <source>Multimed. Tools Appl.</source> <volume>78</volume> (<issue>21</issue>), <fpage>29953</fpage>&#x2013;<lpage>29970</lpage>. <pub-id pub-id-type="doi">10.1007/s11042-018-6748-0</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s11042-018-6748-0">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Gesture+Recognition+Based+on+Skeletonization+Algorithm+and+CNN+with+ASL+Database&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Tan</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2021d</year>). <article-title>Semantic Segmentation for Multiscale Target Based on Object Recognition Using the Improved Faster-RCNN Model</article-title>. <source>Future Gener. Comput. Syst.</source> <volume>123</volume>, <fpage>94</fpage>&#x2013;<lpage>104</lpage>. <pub-id pub-id-type="doi">10.1016/j.future.2021.04.019</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.future.2021.04.019">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Semantic+Segmentation+for+Multiscale+Target+Based+on+Object+Recognition+Using+the+Improved+Faster-RCNN+Model&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Zheng</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2019a</year>). <article-title>Gesture Recognition Based on Binocular Vision</article-title>. <source>Clust. Comput.</source> <volume>22</volume> (<issue>Suppl. 6</issue>), <fpage>13261</fpage>&#x2013;<lpage>13271</lpage>. <pub-id pub-id-type="doi">10.1007/s10586-018-1844-5</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s10586-018-1844-5">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Gesture+Recognition+Based+on+Binocular+Vision&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2019b</year>). <article-title>Gesture Recognition Based on Modified Adaptive Orthogonal Matching Pursuit Algorithm</article-title>. <source>Clust. Comput.</source> <volume>22</volume> (<issue>Suppl. 1</issue>), <fpage>503</fpage>&#x2013;<lpage>512</lpage>. <pub-id pub-id-type="doi">10.1007/s10586-017-1231-7</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s10586-017-1231-7">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Gesture+Recognition+Based+on+Modified+Adaptive+Orthogonal+Matching+Pursuit+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Surface EMG Data Aggregation Processing for Intelligent Prosthetic Action Recognition</article-title>. <source>Neural Comput. Applic</source> <volume>32</volume> (<issue>22</issue>), <fpage>16795</fpage>&#x2013;<lpage>16806</lpage>. <pub-id pub-id-type="doi">10.1007/s00521-018-3909-z</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s00521-018-3909-z">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Surface+EMG+Data+Aggregation+Processing+for+Intelligent+Prosthetic+Action+Recognition&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Manogaran</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2019c</year>). <article-title>Human Lesion Detection Method Based on Image Information and Brain Signal</article-title>. <source>IEEE Access</source> <volume>7</volume>, <fpage>11533</fpage>&#x2013;<lpage>11542</lpage>. <pub-id pub-id-type="doi">10.1109/access.2019.2891749</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/access.2019.2891749">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Human+Lesion+Detection+Method+Based+on+Image+Information+and+Brain+Signal&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Ju</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2019a</year>). <article-title>A Novel Feature Extraction Method for Machine Learning Based on Surface Electromyography from Healthy Brain</article-title>. <source>Neural Comput. Applic</source> <volume>31</volume> (<issue>12</issue>), <fpage>9013</fpage>&#x2013;<lpage>9022</lpage>. <pub-id pub-id-type="doi">10.1007/s00521-019-04147-3</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s00521-019-04147-3">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=A+Novel+Feature+Extraction+Method+for+Machine+Learning+Based+on+Surface+Electromyography+from+Healthy+Brain&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Xiao</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>An Inverse Kinematics Method for Robots after Geometric Parameters Compensation</article-title>. <source>Mech. Mach. Theory</source> <volume>174</volume>, <fpage>104903</fpage>. <pub-id pub-id-type="doi">10.1016/j.mechmachtheory.2022.104903</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.mechmachtheory.2022.104903">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=An+Inverse+Kinematics+Method+for+Robots+after+Geometric+Parameters+Compensation&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Long</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Ship Target Detection and Recognition Method on Sea Surface Based on Multi-Level Hybrid Network[J]</article-title>. <source>J. Beijing Inst. Technol.</source> <volume>30</volume>, <fpage>1</fpage>&#x2013;<lpage>10</lpage>. <pub-id pub-id-type="doi">10.15918/j.jbit1004-0579.20141</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.15918/j.jbit1004-0579.20141">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Ship+Target+Detection+and+Recognition+Method+on+Sea+Surface+Based+on+Multi-Level+Hybrid+Network[J]&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liao</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Multi-object Intergroup Gesture Recognition Combined with Fusion Feature and KNN Algorithm</article-title>. <source>Ifs</source> <volume>38</volume> (<issue>3</issue>), <fpage>2725</fpage>&#x2013;<lpage>2735</lpage>. <pub-id pub-id-type="doi">10.3233/jifs-179558</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3233/jifs-179558">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Multi-object+Intergroup+Gesture+Recognition+Combined+with+Fusion+Feature+and+KNN+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liao</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Occlusion Gesture Recognition Based on Improved SSD</article-title>. <source>Concurrency Comput. Pract. Exp.</source> <volume>33</volume> (<issue>6</issue>), <fpage>e6063</fpage>. <pub-id pub-id-type="doi">10.1002/cpe.6063</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/cpe.6063">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Occlusion+Gesture+Recognition+Based+on+Improved+SSD&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2021a</year>). <article-title>Insulator Faults Detection in Aerial Images from High-Voltage Transmission Lines Based on Deep Learning Model</article-title>. <source>Appl. Sci.</source> <volume>11</volume> (<issue>10</issue>), <fpage>4647</fpage>. <pub-id pub-id-type="doi">10.3390/app11104647</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3390/app11104647">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Insulator+Faults+Detection+in+Aerial+Images+from+High-Voltage+Transmission+Lines+Based+on+Deep+Learning+Model&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2022a</year>). <article-title>Genetic Algorithm-Based Trajectory Optimization for Digital Twin Robots</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>9</volume>, <fpage>793782</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2021.793782</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35083202/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2021.793782">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Genetic+Algorithm-Based+Trajectory+Optimization+for+Digital+Twin+Robots&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>J</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>S</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Self-tuning Control of Manipulator Positioning Based on Fuzzy PID and PSO Algorithm</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>9</volume>, <fpage>817723</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2021.817723</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35223822/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2021.817723">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Self-tuning+Control+of+Manipulator+Positioning+Based+on+Fuzzy+PID+and+PSO+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Duan</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group> (<year>2021b</year>). <article-title>Dynamic Gesture Recognition Algorithm Based on 3D Convolutional Neural Network</article-title>. <source>Comput. Intell. Neurosci.</source> <volume>2021</volume>, <fpage>4828102</fpage>. <pub-id pub-id-type="doi">10.1155/2021/4828102</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/34447430/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1155/2021/4828102">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Dynamic+Gesture+Recognition+Algorithm+Based+on+3D+Convolutional+Neural+Network&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Qi</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2022b</year>). <article-title>Grasping Posture of Humanoid Manipulator Based on Target Shape Analysis and Force Closure</article-title>. <source>Alexandria Eng. J.</source> <volume>61</volume> (<issue>5</issue>), <fpage>3959</fpage>&#x2013;<lpage>3969</lpage>. <pub-id pub-id-type="doi">10.1016/j.aej.2021.09.017</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.aej.2021.09.017">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Grasping+Posture+of+Humanoid+Manipulator+Based+on+Target+Shape+Analysis+and+Force+Closure&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2021c</year>). <article-title>Wrist Angle Prediction under Different Loads Based on GA&#x2010;ELM Neural Network and Surface Electromyography</article-title>. <source>Concurrency Comput.</source> <volume>34</volume>, <fpage>2021</fpage>. <pub-id pub-id-type="doi">10.1002/CPE.6574</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/CPE.6574">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Wrist+Angle+Prediction+under+Different+Loads+Based+on+GA&#x2010;ELM+Neural+Network+and+Surface+Electromyography&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Xiao</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2021d</year>). <article-title>Manipulator Trajectory Planning Based on Work Subspace Division</article-title>. <source>Concurrency Comput.</source> <volume>34</volume>. <pub-id pub-id-type="doi">10.1002/CPE.6710</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/CPE.6710">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Manipulator+Trajectory+Planning+Based+on+Work+Subspace+Division&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2021e</year>). <article-title>Target Localization in Local Dense Mapping Using RGBD SLAM and Object Detection</article-title>. <source>Concurrency Comput. Pract. Exp</source>. <pub-id pub-id-type="doi">10.1002/CPE.6655</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/CPE.6655">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Target+Localization+in+Local+Dense+Mapping+Using+RGBD+SLAM+and+Object+Detection&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Luo</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Ju</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Decomposition Algorithm for Depth Image of Human Health Posture Based on Brain Health</article-title>. <source>Neural Comput. Applic</source> <volume>32</volume> (<issue>10</issue>), <fpage>6327</fpage>&#x2013;<lpage>6342</lpage>. <pub-id pub-id-type="doi">10.1007/s00521-019-04141-9</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s00521-019-04141-9">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Decomposition+Algorithm+for+Depth+Image+of+Human+Health+Posture+Based+on+Brain+Health&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ma</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Grasping Force Prediction Based on sEMG Signals</article-title>. <source>Alexandria Eng. J.</source> <volume>59</volume> (<issue>3</issue>), <fpage>1135</fpage>&#x2013;<lpage>1147</lpage>. <pub-id pub-id-type="doi">10.1016/j.aej.2020.01.007</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.aej.2020.01.007">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Grasping+Force+Prediction+Based+on+sEMG+Signals&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Meng</surname>
<given-names>F.-j.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.-q.</given-names>
</name>
<name>
<surname>Shao</surname>
<given-names>F.-m.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>X.-d.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Fast-armored Target Detection Based on Multi-Scale Representation and Guided Anchor</article-title>. <source>Def. Technol.</source> <volume>16</volume> (<issue>4</issue>), <fpage>922</fpage>&#x2013;<lpage>932</lpage>. <pub-id pub-id-type="doi">10.1016/j.dt.2019.11.009</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.dt.2019.11.009">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Fast-armored+Target+Detection+Based+on+Multi-Scale+Representation+and+Guided+Anchor&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B44">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ming-Ming</surname>
<given-names>L. I.</given-names>
</name>
<name>
<surname>Lei</surname>
<given-names>J. Y.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>C. J.</given-names>
</name>
</person-group> (<year>2019</year>). <source>Multi-target Detection under Road Scenes Based on Video</source>. <publisher-name>Computer Engineering &#x26; Software</publisher-name>. <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Multi-target+Detection+under+Road+Scenes+Based+on+Video&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pan</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Pang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>A New Image Recognition and Classification Method Combining Transfer Learning Algorithm and MobileNet Model for Welding Defects</article-title>. <source>IEEE Access</source> <volume>8</volume>, <fpage>119951</fpage>&#x2013;<lpage>119960</lpage>. <pub-id pub-id-type="doi">10.1109/access.2020.3005450</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/access.2020.3005450">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=A+New+Image+Recognition+and+Classification+Method+Combining+Transfer+Learning+Algorithm+and+MobileNet+Model+for+Welding+Defects&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qiu</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Lu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Gong</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Research on General Detection Method of Coastline and Sea-Sky Line in FLIR Image</article-title>. <source>Acta Armamentarii</source> <volume>40</volume> (<issue>6</issue>), <fpage>1171</fpage>&#x2013;<lpage>1178</lpage>. <pub-id pub-id-type="doi">10.3969/j.issn.1000-1093.2019.06.007</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3969/j.issn.1000-1093.2019.06.007">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Research+on+General+Detection+Method+of+Coastline+and+Sea-Sky+Line+in+FLIR+Image&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ren</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Bao</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>A Review on Human-Computer Interaction and Intelligent Robots</article-title>. <source>Int. J. Info. Tech. Dec. Mak.</source> <volume>19</volume> (<issue>01</issue>), <fpage>5</fpage>&#x2013;<lpage>47</lpage>. <pub-id pub-id-type="doi">10.1142/s0219622019300052</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1142/s0219622019300052">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=A+Review+on+Human-Computer+Interaction+and+Intelligent+Robots&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B48">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ren</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Girshick</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Faster R-CNN: Towards Real-Time Object Detection with Region Proposal Networks</article-title>. <source>IEEE Trans. Pattern Anal. Mach. Intell.</source> <volume>39</volume> (<issue>6</issue>), <fpage>1137</fpage>&#x2013;<lpage>1149</lpage>. <pub-id pub-id-type="doi">10.1109/tpami.2016.2577031</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/27295650/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/tpami.2016.2577031">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Faster+R-CNN:+Towards+Real-Time+Object+Detection+with+Region+Proposal+Networks&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B49">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sandler</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Howard</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Zhmoginov</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>L. C.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Mobilenetv2: Inverted Residuals and Linear Bottlenecks</article-title>. <source>Proc. IEEE Conf. Comput. Vis. pattern Recognit.</source>, <fpage>4510</fpage>&#x2013;<lpage>4520</lpage>. <pub-id pub-id-type="doi">10.1109/cvpr.2018.00474</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/cvpr.2018.00474">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Mobilenetv2:+Inverted+Residuals+and+Linear+Bottlenecks&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B50">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shelhamer</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Long</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Darrell</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Fully Convolutional Networks for Semantic Segmentation</article-title>. <source>IEEE Trans. Pattern Anal. Mach. Intell.</source> <volume>39</volume> (<issue>6</issue>), <fpage>640</fpage>&#x2013;<lpage>651</lpage>. <pub-id pub-id-type="doi">10.1109/TPAMI.2016.2572683</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/27244717/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/TPAMI.2016.2572683">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Fully+Convolutional+Networks+for+Semantic+Segmentation&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B51">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shi</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Path Planning Optimization of Intelligent Vehicle Based on Improved Genetic and Ant Colony Hybrid Algorithm</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>10</volume>, <fpage>905983</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.905983</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35845413/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.905983">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Path+Planning+Optimization+of+Intelligent+Vehicle+Based+on+Improved+Genetic+and+Ant+Colony+Hybrid+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B52">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Yuan</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Bai</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2022a</year>). <article-title>Multi-objective Optimization Design of Ladle Refractory Lining Based on Genetic Algorithm</article-title>. <source>Front. Bioeng. Biotechnol</source> <volume>10</volume> , <fpage>900655</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.900655</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35782507/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.900655">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Multi-objective+Optimization+Design+of+Ladle+Refractory+Lining+Based+on+Genetic+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B53">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Numerical Simulation of Thermal Insulation and Longevity Performance in New Lightweight Ladle</article-title>. <source>Concurr. Comput. Pract. Exper</source> <volume>32</volume> (<issue>22</issue>), <fpage>e5830</fpage>. <pub-id pub-id-type="doi">10.1002/cpe.5830</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/cpe.5830">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Numerical+Simulation+of+Thermal+Insulation+and+Longevity+Performance+in+New+Lightweight+Ladle&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B54">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Kong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Intelligent Human Computer Interaction Based on Non Redundant EMG Signal</article-title>. <source>Alexandria Eng. J.</source> <volume>59</volume> (<issue>3</issue>), <fpage>1149</fpage>&#x2013;<lpage>1157</lpage>. <pub-id pub-id-type="doi">10.1016/j.aej.2020.01.015</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.aej.2020.01.015">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Intelligent+Human+Computer+Interaction+Based+on+Non+Redundant+EMG+Signal&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B55">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Hao</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Multiscale Generative Adversarial Network for Real&#x2010;world Super&#x2010;resolution</article-title>. <source>Concurr. Comput. Pract. Exper</source> <volume>33</volume> (<issue>21</issue>), <fpage>e6430</fpage>. <pub-id pub-id-type="doi">10.1002/CPE.6430</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/CPE.6430">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Multiscale+Generative+Adversarial+Network+for+Real&#x2010;world+Super&#x2010;resolution&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B56">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2022b</year>). <article-title>Low-illumination Image Enhancement Algorithm Based on Improved Multi-Scale Retinex and ABC Algorithm Optimization</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>10</volume>, <fpage>865820</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.865820</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35480971/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.865820">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Low-illumination+Image+Enhancement+Algorithm+Based+on+Improved+Multi-Scale+Retinex+and+ABC+Algorithm+Optimization&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B57">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tan</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Research on Gesture Recognition of Smart Data Fusion Features in the IoT</article-title>. <source>Neural Comput. Applic</source> <volume>32</volume> (<issue>22</issue>), <fpage>16917</fpage>&#x2013;<lpage>16929</lpage>. <pub-id pub-id-type="doi">10.1007/s00521-019-04023-0</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s00521-019-04023-0">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Research+on+Gesture+Recognition+of+Smart+Data+Fusion+Features+in+the+IoT&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B58">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2022a</year>). <article-title>3D Reconstruction Based on Photoelastic Fringes</article-title>. <source>Concurr. Comput. Pract. Exper</source> <volume>34</volume> (<issue>1</issue>), <fpage>e6481</fpage>. <pub-id pub-id-type="doi">10.1002/CPE.6481</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/CPE.6481">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=3D+Reconstruction+Based+on+Photoelastic+Fringes&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B59">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Qian</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Qian</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>F.</given-names>
</name>
<etal/>
</person-group> (<year>2022b</year>). <article-title>Photoelastic Stress Field Recovery Using Deep Convolutional Neural Network</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>10</volume> <fpage>818112</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.818112</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35387296/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.818112">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Photoelastic+Stress+Field+Recovery+Using+Deep+Convolutional+Neural+Network&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B60">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tian</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Gesture Recognition Based on Multilevel Multimodal Feature Fusion</article-title>. <source>Ifs</source> <volume>38</volume> (<issue>3</issue>), <fpage>2539</fpage>&#x2013;<lpage>2550</lpage>. <pub-id pub-id-type="doi">10.3233/jifs-179541</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3233/jifs-179541">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Gesture+Recognition+Based+on+Multilevel+Multimodal+Feature+Fusion&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B61">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Improved Multi-Stream Convolutional Block Attention Module for sEMG-Based Gesture Recognition</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>10</volume>, <fpage>909023</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.909023</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35747495/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.909023">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Improved+Multi-Stream+Convolutional+Block+Attention+Module+for+sEMG-Based+Gesture+Recognition&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B62">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Weng</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Enhancement of Real-Time Grasp Detection by Cascaded Deep Convolutional Neural Networks</article-title>. <source>Concurrency Comput. Pract. Exp.</source> <volume>33</volume> (<issue>5</issue>), <fpage>e5976</fpage>. <pub-id pub-id-type="doi">10.1002/cpe.5976</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1002/cpe.5976">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Enhancement+of+Real-Time+Grasp+Detection+by+Cascaded+Deep+Convolutional+Neural+Networks&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B63">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Attitude Stabilization Control of Autonomous Underwater Vehicle Based on Decoupling Algorithm and PSO-ADRC</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>28</volume>, <fpage>843020</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.843020</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.843020">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Attitude+Stabilization+Control+of+Autonomous+Underwater+Vehicle+Based+on+Decoupling+Algorithm+and+PSO-ADRC&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B64">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xiao</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>An Effective and Unified Method to Derive the Inverse Kinematics Formulas of General Six-DOF Manipulator with Simple Geometry</article-title>. <source>Mech. Mach. Theory</source> <volume>159</volume>, <fpage>104265</fpage>. <pub-id pub-id-type="doi">10.1016/j.mechmachtheory.2021.104265</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.mechmachtheory.2021.104265">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=An+Effective+and+Unified+Method+to+Derive+the+Inverse+Kinematics+Formulas+of+General+Six-DOF+Manipulator+with+Simple+Geometry&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B65">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Genetic-Based Optimization of 3D Burch-Schneider Cage with Functionally Graded Lattice Material</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>10</volume>, <fpage>819005</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.819005</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35155392/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.819005">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Genetic-Based+Optimization+of+3D+Burch-Schneider+Cage+with+Functionally+Graded+Lattice+Material&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B66">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Dynamic Gesture Recognition Using Surface EMG Signals Based on Multi-Stream Residual Network</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>9</volume>, <fpage>779353</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2021.779353</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/34746114/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2021.779353">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Dynamic+Gesture+Recognition+Using+Surface+EMG+Signals+Based+on+Multi-Stream+Residual+Network&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B67">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Hand Medical Monitoring System Based on Machine Learning and Optimal EMG Feature Set</article-title>. <source>Pers. Ubiquit Comput</source>. <pub-id pub-id-type="doi">10.1007/s00779-019-01285-2</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1007/s00779-019-01285-2">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Hand+Medical+Monitoring+System+Based+on+Machine+Learning+and+Optimal+EMG+Feature+Set&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B68">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Zeng</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>H.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Application of PSO-RBF Neural Network in Gesture Recognition of Continuous Surface EMG Signals</article-title>. <source>Ifs</source> <volume>38</volume> (<issue>3</issue>), <fpage>2469</fpage>&#x2013;<lpage>2480</lpage>. <pub-id pub-id-type="doi">10.3233/jifs-179535</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3233/jifs-179535">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Application+of+PSO-RBF+Neural+Network+in+Gesture+Recognition+of+Continuous+Surface+EMG+Signals&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B69">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Self-adjusting Force/Bit Blending Control Based on Quantitative Factor-Scale Factor Fuzzy-PID Bit Control</article-title>. <source>Alexandria Eng. J.</source> <volume>61</volume> (<issue>6</issue>), <fpage>4389</fpage>&#x2013;<lpage>4397</lpage>. <pub-id pub-id-type="doi">10.1016/j.aej.2021.09.067</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.aej.2021.09.067">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Self-adjusting+Force/Bit+Blending+Control+Based+on+Quantitative+Factor-Scale+Factor+Fuzzy-PID+Bit+Control&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B70">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Research on Dual Mode Target Detection Algorithm for Embedded Platform</article-title>. <source>Complexity</source> <volume>2021</volume>, <fpage>1</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1155/2021/9935621</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1155/2021/9935621">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Research+on+Dual+Mode+Target+Detection+Algorithm+for+Embedded+Platform&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B71">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Xiao</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Time Optimal Trajectory Planing Based on Improved Sparrow Search Algorithm</article-title>. <source>Front. Bioeng. Biotechnol.</source> <volume>10</volume>, <fpage>852408</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.852408</pub-id> <ext-link ext-link-type="uri" xlink:href="https://pubmed.ncbi.nlm.nih.gov/35392405/">PubMed Abstract</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.852408">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=Time+Optimal+Trajectory+Planing+Based+on+Improved+Sparrow+Search+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B72">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhao</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>A Tandem Robotic Arm Inverse Kinematic Solution Based on an Improved Particle Swarm Algorithm</article-title>. <source>Front. Bioeng. Biotechnol.</source>, <fpage>2022</fpage>. <pub-id pub-id-type="doi">10.3389/fbioe.2022.832829</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbioe.2022.832829">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=A+Tandem+Robotic+Arm+Inverse+Kinematic+Solution+Based+on+an+Improved+Particle+Swarm+Algorithm&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
<ref id="B73">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhao</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yan</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>A Lightweight Object Detection Network for Real-Time Detection of Driver Handheld Call on Embedded Devices</article-title>. <source>Comput. Intell. Neurosci.</source> <volume>2020</volume>, <fpage>1</fpage>&#x2013;<lpage>12</lpage>. <pub-id pub-id-type="doi">10.1155/2020/6616584</pub-id> <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1155/2020/6616584">CrossRef Full Text</ext-link> &#x7c; <ext-link ext-link-type="uri" xlink:href="https://scholar.google.com/scholar?hl=en&#x0026;as_sdt=0%2C5&#x0026;q=A+Lightweight+Object+Detection+Network+for+Real-Time+Detection+of+Driver+Handheld+Call+on+Embedded+Devices&#x0026;btnG=">Google Scholar</ext-link>
</citation>
</ref>
</ref-list>
</back>
</article>