<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="review-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Physiol.</journal-id>
<journal-title>Frontiers in Physiology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Physiol.</abbrev-journal-title>
<issn pub-type="epub">1664-042X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1536542</article-id>
<article-id pub-id-type="doi">10.3389/fphys.2025.1536542</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Physiology</subject>
<subj-group>
<subject>Review</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>A review of lightweight convolutional neural networks for ultrasound signal classification</article-title>
<alt-title alt-title-type="left-running-head">Zhang et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fphys.2025.1536542">10.3389/fphys.2025.1536542</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Zhang</surname>
<given-names>Bokun</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="author-notes" rid="fn1">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2853411/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Li</surname>
<given-names>Zhengping</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="author-notes" rid="fn1">
<sup>&#x2020;</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Hao</surname>
<given-names>Yuwen</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Wang</surname>
<given-names>Lijun</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Li</surname>
<given-names>Xiaoxue</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Yao</surname>
<given-names>Yuan</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>School of Information Science and Technology</institution>, <institution>North China University of Technology</institution>, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Disaster Medicine Research Center</institution>, <institution>Medical Innovation Research Division of the Chinese PLA General Hospital Beijing</institution>, <institution>China Beijing Key Laboratory of Disaster Medicine</institution>, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Hangzhou Institute of Technology</institution>, <institution>Xidian University</institution>, <addr-line>Xi&#x2019;an</addr-line>, <country>China</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Emergency Department</institution>, <institution>903rd Hospital of PLA Joint Logistic Support Force</institution>, <addr-line>Hangzhou</addr-line>, <country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/77692/overview">Minkyu Ahn</ext-link>, Handing Global University, Republic of Korea</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2134261/overview">Sarada Prasad Dakua</ext-link>, Hamad Medical Corporation, Qatar</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/3004828/overview">Marina Adriana Mercioni</ext-link>, Politehnica University of Timi&#x219;oara, Romania</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Lijun Wang, <email>ljwang@outlook.com</email>; Xiaoxue Li, <email>lixiaoxue@301hospital.com.cn</email>
</corresp>
<fn fn-type="equal" id="fn1">
<label>
<sup>&#x2020;</sup>
</label>
<p>These authors share first authorship</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>24</day>
<month>04</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>16</volume>
<elocation-id>1536542</elocation-id>
<history>
<date date-type="received">
<day>06</day>
<month>12</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>31</day>
<month>03</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2025 Zhang, Li, Hao, Wang, Li and Yao.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Zhang, Li, Hao, Wang, Li and Yao</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Ultrasound signal processing plays an important role in medical image analysis. Embedded ultrasonography systems with low power consumption and high portability are suitable for disaster rescue, but due to the difficulty of ultrasonic signal recognition, operators need to have strong professional knowledge, and it is not easy to deploy ultrasonography systems in areas with relatively weak infrastructures. In recent years, with the continuous development in the field of deep learning and artificial intelligence, lightweight convolutional neural networks have brought new opportunities for ultrasound signal processing. This paper focuses on investigating lightweight convolutional neural networks applied to ultrasound signal classification. Combined with the characteristics of ultrasound signals, this paper provides a detailed review of lightweight algorithms from two perspectives: model compression and operational optimization. Among them, model compression deals with the overall framework to reduce network redundancy, and the latter aims at the lightweight design of the basic operational module &#x201c;convolution&#x201d; in the network. The experimental results of some classical models and algorithms on the ImageNet dataset are summarized. Through the comprehensive analysis, we present some problems and provide an outlook on the future development of lightweight techniques for ultrasound signal classification.</p>
</abstract>
<kwd-group>
<kwd>ultrasound</kwd>
<kwd>signal classification</kwd>
<kwd>lightweight technology</kwd>
<kwd>model compression</kwd>
<kwd>optimization of lightweight network</kwd>
<kwd>convolutional neural network</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Computational Physiology and Medicine</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>Ultrasound imaging is a crucial medical imaging technology that, compared to CT and X-ray, offers portability, simplicity, and no ionizing radiation, making it ideal for deployment in resource-limited environments such as disaster relief. However, ultrasound images are often complex and susceptible to noise, requiring doctors to rely on subjective experience for diagnosis. Integrating artificial intelligence can assist in recognition, but traditional deep convolutional neural networks (<xref ref-type="bibr" rid="B41">Krizhevsky et al., 2012</xref>; <xref ref-type="bibr" rid="B72">Simonyan and Zisserman, 2014</xref>; <xref ref-type="bibr" rid="B75">Szegedy et al., 2015</xref>; <xref ref-type="bibr" rid="B26">He et al., 2016</xref>), with high computational demands and large parameter sizes, are unsuitable for portable ultrasound devices. Lightweight models reduce computational requirements, enabling real-time ultrasound image processing to help doctors diagnose conditions more quickly and accurately, reducing patient wait times (<xref ref-type="bibr" rid="B50">Liu et al., 2023</xref>). Additionally, lightweight models minimize dependence on specialized skills and complex equipment, improving the accessibility and portability of ultrasound-based diagnostics. Given its lower cost compared to CT and X-ray, ultrasound imaging facilitates the widespread adoption of intelligent diagnostic technologies, particularly in developing countries and remote areas (<xref ref-type="bibr" rid="B21">Goo&#xdf;en et al., 2019</xref>). Therefore, the popularization of lightweight technology in the ultrasound examination process can promote intelligent progress in medical image analysis, which can provide support for medical diagnosis and has significant application value and social significance.</p>
<p>In 2016, the first lightweight model SqueezeNet (<xref ref-type="bibr" rid="B35">Iandola, 2016</xref>) was made public to achieve results approximating AlexNet on the ImageNet dataset but with 1/50th of the model size of AlexNet. Song Han and Hinton proposed weight pruning (<xref ref-type="bibr" rid="B25">Han et al., 2015</xref>) and knowledge distillation (<xref ref-type="bibr" rid="B28">Hinton, 2015</xref>) respectively to reduce redundancy in deep network structure from the perspective of model compression. Lightweight techniques, mainly lightweight model construction and model compression, have triggered a large number of influential research and breakthroughs (<xref ref-type="bibr" rid="B9">Chen et al., 2023</xref>).</p>
<p>In the process of developing traditional neural networks, there may be a large amount of redundancy as the depth of the network deepens. This redundancy mainly consists of computational complexity and many parameters. To lighten the network and remove the redundancy, we need to optimize the model itself and the underlying framework, of which the basic modular unit of the underlying framework is &#x201c;convolution&#x201d;. Therefore, we divide the lightweight techniques into two directions: model compression and computational optimization. The former is to compress a large neural network into a lightweight network, mainly including network pruning, knowledge distillation and low-rank decomposition. The latter is to lighten the design of &#x201c;convolution&#x201d;, which is an operational module in the network. The lightweight technique can realize efficient signal analysis in resource-constrained environments. The classification of lightweight technologies is shown in <xref ref-type="fig" rid="F1">Figure 1</xref>.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Classification of lightweight methods.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g001.tif"/>
</fig>
<p>Artificial intelligence and computer-aided diagnostic solutions can significantly standardize medical practice, reduce training time, and improve the quality of ultrasound signals. There are four main research areas.<list list-type="simple">
<list-item>
<p>1. The invocation of machine learning techniques is expected to significantly improve signal quality and make imaging clearer (<xref ref-type="bibr" rid="B14">Deffieux et al., 2021</xref>). Emerging methods such as beamforming, super-resolution and data enhancement techniques have achieved some results (<xref ref-type="bibr" rid="B58">Micucci and Iula, 2022</xref>), but they often require hardware tuning. Despite the difficulty of implementation, research has gradually overcome the limitations of traditional image reconstruction algorithms, especially in the translation of ultrasound physical measurements into visualized images.</p>
</list-item>
<list-item>
<p>2. Artificial intelligence algorithms can help healthcare professionals perform a thorough examination, thus helping to reduce the learning curve of ultrasound scanning (<xref ref-type="bibr" rid="B59">Mischi et al., 2020</xref>)</p>
</list-item>
<list-item>
<p>3. The most competitive solutions currently available are deep learning-based image processing methods compared to traditional feature engineering methods (<xref ref-type="bibr" rid="B70">Sashidhar et al., 2021</xref>). These algorithms show significant advantages in measurement, quantification and computer-aided detection.</p>
</list-item>
<list-item>
<p>4. The application of computer-aided techniques in diagnosis and triage is receiving research attention because these methods can effectively reduce the burden on physicians and improve their efficiency (<xref ref-type="bibr" rid="B80">Van Sloun et al., 2019</xref>).</p>
</list-item>
</list>
</p>
<p>Ultrasound imaging quality is heavily operator-dependent, posing challenges for inexperienced practitioners. Over-filtering and improper gain adjustments, while improving texture smoothness, often reduce the clinical utility of images. Research into computer-aided scanning techniques aims to enhance automation, making ultrasound acquisition more efficient and accessible.</p>
<p>For developing countries, the impact of computer-aided scanning may be even more significant. High-quality ultrasound is inherently costly, many areas with poor infrastructure are not equipped for ultrasound (<xref ref-type="bibr" rid="B21">Goo&#xdf;en et al., 2019</xref>). Also, the lack of experienced sonographers in developing countries prevents patients from undergoing timely ultrasound diagnosis such as prenatal examinations (<xref ref-type="bibr" rid="B59">Mischi et al., 2020</xref>). Therefore, the creation of a computer-aided scanning system could make it possible to perform ultrasound examinations in remote areas with people with only basic anatomical knowledge. Such a system could filter out images of clinical value and send them to radiologists for specialized diagnosis, even if they are thousands of miles away.</p>
<p>Despite its advantages, ultrasound imaging presents unique challenges due to its susceptibility to noise, variable feature scales, and complex temporal characteristics. These signal features not only affect the clarity and accuracy of imaging but also increase the computational burden of deep learning models in recognition and classification. Lightweight convolutional neural networks provide an effective solution for ultrasound signal classification, which achieves efficient signal analysis in resource-limited environments through model compression and computational optimization.</p>
<p>Ultrasound imaging can be divided into A-type, B-type and M-type ultrasound. A-type ultrasound displays the intensity of a single echo, the B-type ultrasound converts the A-type signals into two-dimensional ones for static analysis, and the M-type ultrasound imaging is simpler and does not require complex reconstruction, making it the first choice for portable ultrasound detection equipment. For M-typesignals, periodic signals show stable fluctuations, which is conducive to quantitative evaluation, while non-periodic signals clearly show abnormal features.</p>
<p>As shown in <xref ref-type="fig" rid="F2">Figure 2</xref>, M-ultrasound measures the change of reflection intensity with depth and time along a fixed direction, with the horizontal coordinate indicating time, the vertical coordinate indicating depth information, and the brightness indicating the strength of the reflected signals. M-Ultrasound has a very high temporal resolution and is able to accurately capture the rapid movement of tissues or organs, which is very suitable for detecting dynamic tissues and organs. However, due to this dynamic characteristic, M-Ultrasound is more sensitive to images with scattering noise (<xref ref-type="bibr" rid="B38">Kang et al., 2015</xref>) and acoustic shadow effect (<xref ref-type="bibr" rid="B56">Matsuyama et al., 2022</xref>). Scattering noise is a type of grain-mounted noise in ultrasound imaging due to the coherent superposition of acoustic waves with scattering in tissues, which reduces the contrast and detail resolution of the image. The acoustic impedance difference between the bone and muscle tissues of the human body is extremely large, and after the ultrasound wave propagates through the body to the bone, the ultrasound will undergo total reflection at the bone interface due to the large acoustic impedance difference (<xref ref-type="bibr" rid="B45">Li, 2022</xref>), so the ultrasound wave cannot reach the posterior region. This phenomenon is known as the acoustic shadow effect. <xref ref-type="fig" rid="F3">Figure 3</xref> shows an image of the rib cage under B-mode ultrasound. The randomness and non-Gaussian distribution of the scattering noise make ultrasound denoising an important problem.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>M-type ultrasound image.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g002.tif"/>
</fig>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>B-mode ultrasound rib image.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g003.tif"/>
</fig>
<p>Therefore, for ultrasound signals, if the model focuses too much on higher-order feature correlations might enhance irrelevant noise patterns rather than improving target clarity. <xref ref-type="bibr" rid="B2">Ansari et al. (2023)</xref> annotated the extracted images using the Computer Vision Annotation Tool (CVAT). To enhance the contrast and highlight the liver boundary, the images were also preprocessed using Contrast Limited Adaptive Histogram Equalization (CLAHE) <xref ref-type="bibr" rid="B67">Reza (2004)</xref>. <xref ref-type="bibr" rid="B1">Afsa et al. (2024)</xref> used uses Independent Component Analysis (ICA) to eliminate feature redundancy, extract independent components, and improve computational efficiency and model effect. This method can still maintain high prediction performance in the case of data imbalance. <xref ref-type="bibr" rid="B65">Regaya et al. (2023)</xref> effectively reduced image noise and enhanced feature expression by combining maximal overlap discrete wavelet transform (MODWT) and stochastic resonance (SR) (<xref ref-type="bibr" rid="B13">Dakua et al., 2019</xref>) technology in the preprocessing stage, providing high-quality input data for subsequent cerebral aneurysm segmentation tasks. <xref ref-type="bibr" rid="B3">Ansari et al. (2024)</xref> reviewed in detail preprocessing methods such as data enhancement and denoising for ultrasound signals, which effectively solved data scarcity and image quality problems and provided the possibility of building an end-to-end deep learning system.</p>
<p>The application of lightweight technology greatly reduces the computing resource requirements for ultrasound signal classification, making real-time ultrasound analysis possible. However, due to the high temporal resolution of ultrasound signals (such as M-mode ultrasound), the model is required to quickly process time series data to provide immediate feedback. Therefore, parallel computing technologies, such as GPU parallel processing and field programmable gate arrays (FPGAs), play a key role in real-time ultrasound feedback. <xref ref-type="bibr" rid="B94">Zhai et al. (2019a)</xref> proposed a hardware architecture based on Zynq SoC to accelerate the calculation of Lattice Boltzmann (LB) method (<xref ref-type="bibr" rid="B57">Mazzeo and Coveney, 2008</xref>). LB can be efficiently implemented on a variety of parallel architectures, ranging from general purpose graphics processing units (GPGPU) (<xref ref-type="bibr" rid="B42">Kuznik et al., 2010</xref>) and supercomputers (<xref ref-type="bibr" rid="B16">Djelouat et al., 2018</xref>). Based on the above research, (<xref ref-type="bibr" rid="B95">Zhai et al., 2019b</xref>) optimized the HemeLB model, designed an acceleration solution based on Zynq SoC and GPU, and proposed a real-time visualization framework, providing an efficient, scalable, and user-friendly tool for clinical application of hemodynamic simulation. <xref ref-type="bibr" rid="B18">Esfahani et al. (2020)</xref> proposed an integrated pipeline for cerebral aneurysm blood flow simulation and real-time visualization. This pipeline provides an efficient clinical tool for cerebral aneurysm blood flow simulation and visualization by combining GPU-accelerated HemeLB and a real-time rendering engine.</p>
<p>Through parallel computing, convolution operations, feature extraction, and signal classification can be performed simultaneously, thereby reducing latency and improving diagnostic efficiency. In addition, parallel computing can also be combined with deep learning technologies, such as the self-attention mechanism in the Transformer architecture, to effectively improve the processing capabilities of ultrasound time series signals and provide stronger intelligent auxiliary support for portable ultrasound devices.</p>
</sec>
<sec id="s2">
<title>2 Model compression</title>
<p>Model compression refers to the compression of model volume to achieve similar accuracy as the original model by reducing the model size, removing over-parameterization redundancy and structural redundancy, and reducing the memory footprint. Model compression based on computer-aided diagnostic techniques can further enhance the flexibility and deployment capability of ultrasound diagnostic systems. According to different processing ideas, model compression techniques can be mainly classified into network pruning, knowledge distillation and low-rank decompossion. <xref ref-type="table" rid="T1">Table 1</xref> briefly summarises the characteristics of the three basic types of model compression methods.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>The comparison of basic methods for model compression.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Method</th>
<th align="left">Applicable layers</th>
<th align="left">Description</th>
<th align="left">Advantages and disadvantages</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Network Pruning</td>
<td align="left">Convolutional layer and fully connected layer</td>
<td align="left">Removing non-essential redundancies from the pre-trained model while maintaining accuracy</td>
<td align="left">Enhances generalization and reduces overfitting risk but requires specialized libraries and hardware</td>
</tr>
<tr>
<td align="left">Knowledge Distillation</td>
<td align="left">Convolutional layer and fully connected layer</td>
<td align="left">Allowing the student model to match the teacher model&#x2019;s performance with lower computational and memory costs</td>
<td align="left">Speeds up training and enhances performance but is limited to teacher-student setups</td>
</tr>
<tr>
<td align="left">Low-rank Decomposition</td>
<td align="left">Convolutional layer and fully connected layer</td>
<td align="left">Decomposing convolution kernels to reduce redundancy</td>
<td align="left">Improves computational efficiency; but hard to implement and decomposition operation requires a lot of computation</td>
</tr>
</tbody>
</table>
</table-wrap>
<sec id="s2-1">
<title>2.1 Network pruning</title>
<p>In recent years, network pruning has been widely studied as a technique to reduce the computational and storage requirements of neural networks, especially for the compression of deep networks. Network pruning is used to eliminate non-critical redundancies in the pre-trained model without affecting the accuracy of the model. The process of network pruning is shown in <xref ref-type="fig" rid="F4">Figure 4</xref>, where a scaling factor is assigned to each channel of the convolutional layer (<xref ref-type="bibr" rid="B43">Lai et al., 2018</xref>). During the training process, these scaling factors are constrained by sparse regularization to automatically identify unimportant channels. After constraints, channels with smaller scaling factors are pruned. After pruning, the model structure becomes more compact.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Network pruning.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g004.tif"/>
</fig>
<p>Network pruning techniques are widely used in ultrasound image segmentation and classification tasks, where the computational requirements of the model are reduced by pruning. M-ultrasound dynamically tracks data from only a single scan line, which has less data compared to B-ultrasound and 3D-ultrasound but requires a higher temporal resolution. The network pruning technique removes redundant convolutional kernels or channels and is suitable for M-Ultrasound. The pruned network is able to preserve dynamic signal features while reducing computational requirements (<xref ref-type="bibr" rid="B51">Liu et al., 2017</xref>). Usually, network pruning can be categorized into unstructured pruning and structured pruning by granularity.</p>
<sec id="s2-1-1">
<title>2.1.1 Untructured prunning</title>
<p>Unstructured pruning removes individual weights from the model, which minimizes the number of parameters and is commonly used for fine-grained pruning.</p>
<p>Song Han proposed the concept of Deep Compression in 2016 in order to solve the problem that neural networks are computationally and memory intensive and difficult to deploy on systems with limited hardware resources (<xref ref-type="bibr" rid="B26">He et al., 2016</xref>). It consists of three main stages: pruning, trained quantization, and Huffman coding. This deep compression concept compresses the neural network without compromising accuracy. Experiments on AlexNet (<xref ref-type="bibr" rid="B41">Krizhevsky et al., 2012</xref>), VGG-16 (<xref ref-type="bibr" rid="B72">Simonyan and Zisserman, 2014</xref>) and LeNet (<xref ref-type="bibr" rid="B19">Filters&#xe2;&#x20ac;&#x2122;Importance, 2016</xref>) networks were compressed by a factor of 35, 49, and 39, respectively, without loss of accuracy. After pruning, complex neural networks can be used in mobile applications with limited application size and download bandwidth.</p>
<p>However, magnitude-based weight pruning reduces a large number of parameters in the fully connected layer and may not sufficiently reduce the computational cost of the model due to irregular sparsity in the pruned network. <xref ref-type="bibr" rid="B19">Filters&#xe2;&#x20ac; &#x2122;Importance (2016)</xref> proposed to remove filters that have less impact on the accuracy of the output. Unlike pruning weights, this method does not produce sparse connection patterns.</p>
<p>Network pruning can also be combined with other model compression techniques. <xref ref-type="bibr" rid="B63">Park and No (2022)</xref> proposed a &#x2018;prune-then-distill framework&#x2019;, where the teacher model is first pruned to make it transferable and then distilled into the student model. Experiments have shown that distillation of the pruned teacher model can outperform the unpruned teacher model, which reverses the assumption that unpruned teacher networks are always more effective.</p>
<p>As research progressed, network pruning did not focus just on a single weight and certain modules of the neural network, but aimed at an entire layer of the network. <xref ref-type="bibr" rid="B48">Liao et al. (2023)</xref> investigated an EGP entropy-guided pruning algorithm, which targets layers in the network with low entropy values and prioritizes pruning their connections, eventually removing them completely. Through validation in popular models such as ResNet-18 and Swin-T (<xref ref-type="bibr" rid="B52">Liu et al., 2021</xref>), the EGP algorithm significantly compresses the depth of the model. The study reveals that unstructured pruning can also reduce the model depth.</p>
<p>Unstructured pruning has a promising application in ultrasound signal classification tasks. In particularly, it can significantly improve the computational efficiency and resource adaptability of the model when dealing with MMA image signals and time-domain signals. Experiments have shown that sparse networks can enhance the robustness of small-scale signals and can improve the accuracy of ultrasound signal classification (<xref ref-type="bibr" rid="B73">Srinivas et al., 2017</xref>). Meanwhile, unstructured pruning can be used to reduce unimportant temporal correlation weights, thus strengthening the model&#x2019;s focus on key temporal features (<xref ref-type="bibr" rid="B4">Ayle et al., 2022</xref>), and improving the ability to classify heart rate or valve motion signals. In ultrasound time-domain signals, unstructured pruning can accurately remove low-contributing weights by weighted sparsity constraints in a specific time range, thus more efficiently processing signal parts with several energy species.</p>
<p>Unstructured pruning is promising for research due to its higher pruning rate and its ability to be combined with other model optimization techniques. In the future, unstructured pruning may combine dynamic pruning and sparse training techniques more often, allowing unstructured pruning methods to improve sparsity while enhancing hardware adaptability, thus speeding up inference.</p>
</sec>
<sec id="s2-1-2">
<title>2.1.2 Structured pruning</title>
<p>Structured Pruning removes entire neurons, convolutional kernels, channels, or layers from a neural net. This type of pruning is easier to accelerate and is suitable for standard hardware platforms such as CPUs and GPUs.</p>
<p>In 2018, <xref ref-type="bibr" rid="B92">Yu et al. (2018)</xref> proposed the Neuron Importance Score Propagation (NISP) algorithm, where NISP assigns an importance score to each neuron to measure its overall contribution to network performance. By calculating the sensitivity of the network output to each neuron, pruning decisions can be made at the neuron level. NISP can keep neurons that contribute significantly to the final prediction in each layer, removing the least important neurons in the neural network. For ultrasound signal processing tasks, pruning out redundant channels and neurons can result in loss of high-frequency details in the ultrasound image, affecting edge clarity, especially in small lesion detection tasks (<xref ref-type="bibr" rid="B5">Bria et al., 2020</xref>). In contrast, this structured pruning prioritizes the retention of the extraction layer of key features and reduces the damage to high-frequency details (<xref ref-type="bibr" rid="B82">Wang et al., 2021</xref>).</p>
<p>Structured pruning usually imposes sparse constraints on the weight parameters and prunes some unimportant weights during the training process. <xref ref-type="bibr" rid="B71">Shao and Shin (2022)</xref> proposed a dynamic scheme that imposes sparse constraints based on the filter weights. This method evaluates the structure of the model by its performance in real-time and dynamically prunes it according to the current performance. <xref ref-type="bibr" rid="B84">Wen et al. (2016)</xref> proposed the structural sparsity learning (SLL) method to regularize the filters, channels, and layer depths in neural networks. This approach allows deep neural networks (DNNs) to learn more compact structures without loss of accuracy. The compactness of DNNs speeds up DNN evaluation on CPUs and GPUs using off-the-shelf libraries.</p>
<p>Structured pruning requires modifying the network architecture and implementing complex gradient update rules will offset some of the efficiency gains. For some more complex and deeper network structures, the sparsity at different levels of the network may exhibit different properties, making it a challenge to maintain sparsity without loss of accuracy. <xref ref-type="bibr" rid="B23">Gupta et al. (2024)</xref> proposed a novel, mechanics-inspired structured construction method. Similar to &#x201c;Torque,&#x201d; a force is applied during training to adjust convolutional layer weights around pivot points. This increases weight density near the pivot while promoting sparsity further away, enabling filter pruning with minimal information loss.</p>
<p>Structured pruning simplifies the model structure and improves the storage efficiency of the model by removing redundant convolutional kernels, while enhancing the corresponding ability of the model in the regions where signal changes are obvious, enabling ultrasound image analysis to be realized on portable ultrasound detectors. The structured pruning processed model can focus its performance on capturing the periodic fluctuation characteristics of the time domain signal (<xref ref-type="bibr" rid="B88">Ye et al., 2024</xref>), which significantly improves the sensitivity to signal changes and classification efficiency.</p>
</sec>
</sec>
<sec id="s2-2">
<title>2.2 Knowledge distillation</title>
<p>Knowledge Distillation (<xref ref-type="bibr" rid="B28">Hinton, 2015</xref>) is an another model compression technique whose goal is to transfer knowledge from a larger, better performing &#x2018;teacher model&#x2019; to a smaller &#x2018;student model&#x2019;, thus allowing the student model to achieve performance close to that of the teacher model with fewer computational resources and memory usage.</p>
<p>Knowledge distillation can reduce the size of the model regardless of the structural differences between the teacher and student models. When training the student model, the softmax output probability distribution of the teacher model is used as the training target, and a method is proposed to control the output probability distribution with a &#x2018;temperature&#x2019; parameter, which can make the target &#x2018;soft&#x2019;. Given the logits z of the network, the category probability p of an image is calculated as <xref ref-type="disp-formula" rid="e1">Equation 1</xref>.<disp-formula id="e1">
<mml:math id="m1">
<mml:mrow>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mi>T</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>/</mml:mo>
<mml:mi>T</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>exp</mml:mi>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>/</mml:mo>
<mml:mi>T</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>where T is the temperature parameter. When T &#x3d; 1, the standard softmax function is obtained. As T increases, the probability distribution produced by the softmax function becomes softer, thus providing more information. As shown in <xref ref-type="fig" rid="F5">Figure 5</xref>, knowledge distillation can be categorized into logit-based distillation and feature-based distillation based on the location of the knowledge in the teacher model (<xref ref-type="bibr" rid="B8">Chen et al., 2024</xref>).</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Knowledge distillation.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g005.tif"/>
</fig>
<p>Knowledge distillation can refine important timing features in the ultrasound signal. For example, in heart valve motion signals, the teacher model can extract the key time points of the waveform and guide the student model for feature extraction with low computational complexity. Meanwhile, knowledge distillation can obtain dynamic patterns of ultrasound signal patterns. <xref ref-type="bibr" rid="B66">Ren et al. (2023)</xref> used knowledge distillation to compress the laws of ultrasound signal waveform changes in the teacher model into feature representations in the student model for analyzing real-time cardiovascular signals. It has been shown that a deep denoising model incorporating convolutional neural networks can effectively mitigate such interference while maintaining key features in the signal (<xref ref-type="bibr" rid="B58">Micucci and Iula, 2022</xref>). Through distillation, the student model can learn the noise-resistant properties of the teacher model and improve the robustness of ultrasound signal classification. However, if the teacher model itself is sensitive to noise, it can be combined with adversarial training to enhance the robustness of the student model to noise.</p>
<sec id="s2-2-1">
<title>2.2.1 Logit-based distillation</title>
<p>In logit-based distillation, the student model learns the final logits of the teacher model output layer, which are the global representation. Since only the outputs need to be learned, the student model and the teacher model can have different architectures. Even if the student model is much smaller than the teacher model, it can still get good performance by imitating the output. <xref ref-type="bibr" rid="B68">Romero et al. (2014)</xref> and <xref ref-type="bibr" rid="B60">Mishra and Marr (2017)</xref> both verified that knowledge distillation can effectively improve deep network training, especially for the student model, which is shallow in depth and low in accuracy, and can learn more and more detailed features from the deep teacher model.</p>
<p>The attention mechanism also plays a key role in the development of knowledge distillation. <xref ref-type="bibr" rid="B93">Zagoruyko and Komodakis (2016)</xref> proposed a method for applying attention mechanism in the knowledge distillation process, called attention transfer. This approach extracts attention graphs from specific layers of the teacher model as a bridge for transferring knowledge between the teacher model and the student model, which allows the student model to fully learn the hierarchical information within the model. <xref ref-type="bibr" rid="B37">Jin et al. (2023)</xref> integrated the attention mechanism by aligning logits at the instance, batch and category levels, focusing on different levels of important features, and optimizing the knowledge transfer process. This enables the model to capture detailed features and transfer information more effectively, especially in complex ultrasonic signal processing tasks.</p>
<p>Focusing on &#x201c;Neural Collapse&#x201d; (<xref ref-type="bibr" rid="B62">Papyan et al., 2020</xref>), a phenomenon that refers to a series of geometric patterns that appear when a deep neural network approaches zero training error in an image classification task, Papyan et al. proposed a new perspective for optimizing knowledge distillation. Neural collapse simplifies the teacher-student learning process, allowing smaller student models to capture the key structures of the teacher model more easily.</p>
<p>Current logit-based distillation, because of the conflict between standard distillation loss and cross-entropy loss, leads to incorrect predictions even by highly accurate teacher models. <xref ref-type="bibr" rid="B74">Sun et al. (2024)</xref> introduced &#x201c;refund logit-based distillation&#x201d; to address the limitations of the current logit-based distillation. It is also effective in suppressing overfitting and eliminating potential misinformation from the teacher model while maintaining class relevance, ultimately allowing the student model to gain more valuable knowledge.</p>
<p>
<xref ref-type="bibr" rid="B98">Zhao et al. (2022)</xref> proposed an improved knowledge distillation algorithm, Decoupled Knowledge Distillation (DKD), which decouples the loss of distillation into two parts: target category knowledge distillation (TCKD) and non-target category knowledge distillation (NCKD). The former focuses on the target categories and conveys the prediction results of the teacher model for the correct categories; the latter focuses on the probability distribution among the non-target categories and preserves the inter-class relationships.</p>
<p>Logits-based distillation enables the student model to learn the discriminative method of fuzzy samples more accurately by transferring the teacher model&#x2019;s confidence difference in the classified samples, which can realize the effective recognition of weak features (<xref ref-type="bibr" rid="B98">Zhao et al., 2022</xref>) in ultrasound image signals, such as low-contrast lesions. In multitask classification, logit-based distillation can convey the recognition ability of the teacher model for complex signals and improve the adaptability of the student model for dynamic signals.</p>
</sec>
<sec id="s2-2-2">
<title>2.2.2 Feature-based distillation</title>
<p>Feature-based distillation learns feature representations of data samples at different levels, focusing on the local perceptual ability of the model and the expressive ability of the middle layer. Feature distillation can outperform logit-based distillation but is relatively complex to implement and requires additional computation and memory consumption to refine deep features during training (<xref ref-type="bibr" rid="B27">Heo et al., 2019</xref>).</p>
<p>Optimization for feature-based distillation often starts with the structure of the teacher-student models. Since the mapping of deep neural network models from the input space to the output space needs to go through many layers, <xref ref-type="bibr" rid="B89">Yim et al. (2017)</xref> defined the knowledge to be transmitted by the information flow of features between layers, which is obtained by calculating the inner product between the features of the two layers. <xref ref-type="bibr" rid="B34">Huang et al. (2023)</xref> used neuron selectivity to align selectivity patterns between teacher and student models, enhancing student network performance. They also introduce a feature-based distillation strategy, including multi-scale feature distillation, which overcomes single-scale limitations, and self-mutual information distillation, combining self-supervision in the student model with mutual supervision from the teacher. (<xref ref-type="bibr" rid="B27">Heo et al., 2019</xref>) placed the distillation location before the ReLU activation function of the neural network, which eliminated the redundant information that would adversely affect model compression and allowed the student network to learn more effective information from the teacher network. <xref ref-type="bibr" rid="B10">Chen et al. (2022)</xref> proposed to improve the feature distillation by using a projector ensemble (projector ensemble) to improve the feature distillation. Adding multiple projectors to the student model solves the mismatch between the teacher and student feature spaces and improves the performance of the image classification task.</p>
<p>
<xref ref-type="bibr" rid="B49">Liu et al. (2024)</xref> proposed a large kernel attention network based on pyramid segmentation, using the dynamic feature distillation module can extract the features of different layers, effectively improved the performance of the image super-resolution model. <xref ref-type="bibr" rid="B79">Tian et al. (2019)</xref> developed a novel distillation technique by using the Contrastive Learning (<xref ref-type="bibr" rid="B44">Le-Khac et al., 2020</xref>) approach to develop a novel distillation technique that enables teacher and student models to project the same inputs onto adjacent representations and different inputs onto separated representations.</p>
<p>Feature-based distillation requires aligning different levels of feature representations between the teacher model and the student model, which leads to the alignment of the two in feature space becoming challenging. Also passing information about the middle level of the teacher model increases the training time and memory requirements of the model. Whereas the training cost of the logit-based distillation is lower, but the performance is not satisfactory compared to feature-based distillation. Therefore, both of them still need to be further optimized to reduce the problems of knowledge distillation in terms of complexity, computational cost and task suitability.</p>
</sec>
</sec>
<sec id="s2-3">
<title>2.3 Low-rank decomposition</title>
<p>Most deep neural networks are over-parameterized and exhibit large computational overhead, making signal recognition inefficient. Low-rank decomposition <xref ref-type="bibr" rid="B40">Kolda and Bader (2009)</xref> refers to sparsifying the convolution kernel matrix by combining dimensions and imposing low-rank constraints. Since the weight vectors are mostly distributed in low-rank subspaces, the convolution kernel matrix can be reconstructed with a small number of basis vectors to achieve the purpose of reducing the storage space (<xref ref-type="bibr" rid="B11">Cheng et al., 2017</xref>). By approximate decomposition of the weight matrix or feature representation of the neural network, redundant information is removed and the model is made more compact. However, over-decomposition may weaken the ability to perceive the dynamic changes in time and have an impact on the temporal consistency of the ultrasound signal. By using an adaptive low-rank approximation method (<xref ref-type="bibr" rid="B53">Lu, 2024</xref>), the decomposition level is dynamically adjusted to avoid the loss of critical time series information.</p>
<p>
<xref ref-type="bibr" rid="B90">Yin et al. (2022)</xref> proposed a budget-aware neural network compression method based on Tucker decomposition (<xref ref-type="bibr" rid="B83">Weber et al., 2025</xref>), called BATUDE. Maintain or improve model performance while meeting computational budget constraints through automated tensor rank selection and globally optimal rank learning strategies. The method not only simplifies the training process, but also enables the model to automatically learn features suitable for specific data, improving the effectiveness of feature extraction while providing a more favourable feature representation for subsequent image recognition tasks. Futhermore, <xref ref-type="bibr" rid="B87">Yadav et al. (2022)</xref> proposed an efficient neighbor search method based on matrix decomposition, which is optimized for cross-encoder models. Through matrix decomposition technology, the computational cost of neighbor search is significantly reduced while maintaining high retrieval accuracy. This method provides a new idea for solving the problem of efficient search in large-scale data sets. Low-rank decomposition can extract low-dimensional structures from high-dimensional data and reveal the interactions between different modalities.</p>
<p>The diversity and complexity of biomedical data require new data analysis methods. Low-rank decomposition, as a powerful model compression technology, is widely used in medical signal processing, such as image signal denoising, super-resolution reconstruction, feature extraction, etc. CP decomposition (<xref ref-type="bibr" rid="B99">Zhou et al., 2019</xref>) and Tucker decomposition are particularly prominent in image reconstruction and noise removal (<xref ref-type="bibr" rid="B85">Wu et al., 2018</xref>). <xref ref-type="bibr" rid="B6">Burch et al. (2025)</xref> systematically reviewed the applications of low-rank factorization in biomedical data analysis and explored the potential of quantum computing to address the challenges faced by traditional low-rank decomposition. Combined with quantum computing, tensor decomposition is expected to further promote the development of precision medicine, especially in terms of data scale and processing efficiency.</p>
</sec>
</sec>
<sec id="s3">
<title>3 Operational optimization</title>
<p>With the increasing demand for ultrasound signal classification tasks in real-time and embedded devices, it is especially important to design optimization strategies that efficiently process ultrasound data and balance real-time and accuracy. The operational optimization technique focuses on considering how to enhance the extraction ability of key features and reduce redundant calculations. Convolution serves as the logical basis for the operation of the model, and the lightweight design of convolution can maximize the computational efficiency of the network, which is convenient for the model to be used on mobile devices such as portable ultrasound detectors.</p>
<p>In this chapter, lightweight optimization algorithms suitable for ultrasound signal classification, including decoupling and portability modules, will be explored in detail in light of the ultrasound signal characteristics, especially for the temporal dynamic characteristics of M-mode ultrasound signals and the frequency domain characteristics of time-domain signals, which will provide technical support for further advancing ultrasound diagnosis.</p>
<sec id="s3-1">
<title>3.1 Decoupling</title>
<p>In convolutional neural networks, there are dependencies between modules, which can be as small as a certain weight or as large as the entire network layer. Therefore, the coupling degree can be utilized to define the dependency of modules in the model, and the lower the coupling degree, the lower the dependency between modules, and the greater the independence, reusability and portability of modules. Classical lightweight modules such as deeply separable convolution (<xref ref-type="bibr" rid="B31">Howard, 2017</xref>) and channel shuffling (<xref ref-type="bibr" rid="B97">Zhang et al., 2018</xref>) can be explained by the idea of decoupling. As shown in <xref ref-type="fig" rid="F6">Figure 6</xref>, this section will divide some manually designed convolutional structures with the concept of decoupling, which can be mainly categorized into two parts: calculation decoupling and tensor decoupling.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Classification of decoupling methods.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g006.tif"/>
</fig>
<sec id="s3-1-1">
<title>3.1.1 Calculation decoupling</title>
<p>The size of the convolution kernel determines the extent of the sensory field on the input image. The larger the convolution kernel, the larger the perception field and the better the feature extraction effect. In order to ensure the classification effect, most early neural network models, such as AlexNet <xref ref-type="bibr" rid="B41">Krizhevsky et al. (2012)</xref>, used large convolution kernels for feature extraction.</p>
<p>However, large convolutional kernels significantly increase computation and memory consumption for processing high-resolution images such as medical images, leading to a decrease in the efficiency of the model during training and testing, at the same time, larger convolutional kernels imply a wider sensory field, which can mishandle important local features in the ultrasound images. Researchers have therefore explored convolution kernel sizing. Replacing a large-size convolutional kernel with multiple small-size convolutional kernels can be referred to as calculation decoupling because it changes the way the original convolution operates. This idea originates from Inception V3 proposed by <xref ref-type="bibr" rid="B75">Szegedy et al. (2015)</xref>. As shown in <xref ref-type="fig" rid="F7">Figure 7</xref>, the 5&#x2a;5 convolution was replaced with a multilayer network with fewer parameters: the first layer is a 3&#x2a;3 convolution, and the second layer is a fully connected layer on the first 3&#x2a;3 output grid. The replacement reduces the number of model parameters by 27.8% and the sensory field is unchanged before and after the split.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>The 5&#x2a;5 convolution is replaced by a smaller one.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g007.tif"/>
</fig>
<p>
<xref ref-type="bibr" rid="B81">Wang et al. (2018)</xref> were inspired by the Inception model and proposed the PeleeNet model. PeleeNet introduces a 2-way dense layer structure, as shown in <xref ref-type="fig" rid="F8">Figure 8</xref>, which is a parallel structure of multi-scale convolutional kernels (e.g., 1&#x2a;1, 3&#x2a;3), and fuses different scales of sensory fields within a single dense layer, which effectively enhances the ability to capture features of different scales. In the eight-neighborhood pixel, the 3&#x2a;3 convolution is the smallest odd convolution kernel size that can capture the features, so various lightweight models often use the 3&#x2a;3 convolution for the convolution splitting operation to reduce the computation amount of the model.</p>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>2-way dense layer in PeleeNet.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g008.tif"/>
</fig>
<p>In addition to using 3&#x2a;3 regular convolutions to reduce the number of network parameters, 1&#x2a;1 convolutions can be used to perform dimension upscaling and dimension downscaling operations on feature maps to achieve the same purpose. Point-by-point convolution (<xref ref-type="bibr" rid="B33">Hua et al., 2018</xref>) computes a linear combination of the output of the depth convolution by 1&#x2a;1 convolution to obtain a new feature.</p>
<p>Although there are differences in the design structure of the above decoupling strategies, their essence is realized by reducing the size of the number of operations in the multiplication process, which mathematically disassembles the standard single multiplication operation and replaces some of the multiplication operations with addition operations to achieve the effect of improving the speed of the model operation.</p>
<p>The DeepShift (<xref ref-type="bibr" rid="B17">Elhoushi et al., 2021</xref>) replaces the floating-point multiplication operation in the original forward propagation by performing shift-by-bit and inverse-by-bit operations to reduce the computation time required in the model inference process. Shift and inverse operations are faster in hardware circuit devices, so using them to replace the product operation can speed up the model. DeepShift&#x2019;s design idea is novel, using the underlying algorithm to thoroughly accelerate the convolution from the perspective of accelerating hardware resources.</p>
<p>At this stage, the model design approach for convolutional operations relies heavily on the design of the underlying hardware devices, and the binary arithmetic process will make the power operation unavoidable errors in the substitution process, which will lead to the model being limited in practical applications. However, this kind of operation substitution can still be used as a direction for future research.</p>
</sec>
<sec id="s3-1-2">
<title>3.1.2 Tensor decoupling</title>
<p>Unlike the calculation decoupling strategy that reduces the amount of computation by replacing the large-size convolution with a small-size convolution, the tensor decoupling operation chooses to disassemble the conventional convolution in the spatial dimension and the channel dimension and performs it in steps, separating the variables or features that are originally closely connected, so that certain parts or parameters in the model can be varied independently, thus reducing the interdependence between them.</p>
<p>The design of conventional depth-separable convolution belongs to the category of tensor decoupling, as shown in <xref ref-type="fig" rid="F9">Figure 9</xref>, depth-separable convolution <xref ref-type="bibr" rid="B31">Howard (2017)</xref> decomposes the conventional 3D convolution into a depth convolution in two-dimensional space and a point-by-point convolution that modifies the number of channels to turn the 3D input features into independent 2D planar features and channel dimensions. Deep convolution extracts local spatial features in each channel and combines them with point-by-point convolution to complete feature fusion. The decoupling in depth-separable convolution reduces computational complexity and number of parameters, making DS Conv particularly suitable for environments with limited computational resources, such as portable ultrasound devices. The depth-separable convolution can also be combined with an adaptive filter (<xref ref-type="bibr" rid="B91">You and Crebbin, 2022</xref>) for removing scattering noise from ultrasound images to improve the quality of ultrasound signals.</p>
<fig id="F9" position="float">
<label>FIGURE 9</label>
<caption>
<p>Depthwise separable convolution. <bold>(a)</bold> Standard Convolution Filters. <bold>(b)</bold> Depthwise Convolution Filters <bold>(c)</bold> Pointwise Convolution Filters.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g009.tif"/>
</fig>
<p>However, traditional design methods do not consider the interrelationships of features within a convolutional kernel, which leads to limited performance of the model in complex feature learning. To address this limitation, <xref ref-type="bibr" rid="B46">Li et al. (2022)</xref> proposed Blueprint Separable Convolution (BSConv), focusing on decoupling and recombination within the convolutional kernel. BSConv separates and recombines features efficiently by introducing a different feature separation strategy, which explicitly takes into account the interactions of features within the convolutional kernel.</p>
<p>The Fire Module in SqueezeNet <xref ref-type="bibr" rid="B35">Iandola (2016)</xref> is also a decoupled calculation, the Fire Module consists of the squeeze layer and the expand layer. The &#x201c;compression-expansion&#x201d; operation in this design decouples spatial features (3&#x2a;3 convolution) and channel features (1&#x2a;1 convolution), allowing the model to separately model inter-channel and spatial domain features without having to use a single standard convolution operation to process all features simultaneously.</p>
<p>While the decoupling operations in the above models all disassemble the N&#x2a;N convolution into a combination of n&#x2a;n convolutions, in InceptionV3 <xref ref-type="bibr" rid="B76">Szegedy et al. (2016)</xref> the authors propose to disassemble a square convolution of one (7&#x2a;7) into a stack of 1&#x2a;7 and 7&#x2a;1. This decoupling operation reduces the 49 multiplication operations to 7.</p>
<p>Features belonging to different channels after decoupling cannot be transferred. The channel shuffling operation in ShuffleNet (<xref ref-type="bibr" rid="B97">Zhang et al., 2018</xref>) disrupts the feature map channels so that features originally belonging to different channels can be mixed together in the subsequent convolution operation, thus realizing the exchange of information between different channels.</p>
</sec>
</sec>
<sec id="s3-2">
<title>3.2 Portability modules</title>
<p>Convolutional modules are artificially constructed application programming interfaces (APIs) that can be invoked directly in programming using abbreviations, and wrapper packages usually contain fixed convolutional structure modules for high portability. In this section, several lightweight portable modules that can be applied to ultrasound image classification are systematically described.</p>
<sec id="s3-2-1">
<title>3.2.1 Residual module</title>
<p>The residual module proposed by <xref ref-type="bibr" rid="B26">He et al. (2016)</xref> is a key part of the advancement of lightweight development of convolutional neural networks. The deeply separable convolution in MobileNetV1 (<xref ref-type="bibr" rid="B31">Howard, 2017</xref>) and the inverse residual in MobileNetV2 (<xref ref-type="bibr" rid="B69">Sandler et al., 2018</xref>) draw on the principle of the residual module. In mathematical statistics, residuals are usually used to represent the difference between the actual observed value and the fitted value. As shown in <xref ref-type="fig" rid="F10">Figure 10</xref>, in the structure of a neural network a stacked layer, when the input is x, the learned feature is the mapping H(x). Compared with the original feature H(x), the difference is easier to learn directly, so the residuals as the difference computes the learned feature can be expressed as: F(x) &#x3d; H(x)-x and the original mapping is reshaped as F(x)&#x2b;x.</p>
<fig id="F10" position="float">
<label>FIGURE 10</label>
<caption>
<p>Basic residual module.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g010.tif"/>
</fig>
<p>From the above equation, when the residual F(x) is 0, the stacking layer only does constant mapping and the network performance should remain unchanged. But in fact, the residuals are not 0, thus causing the stacking layer to constantly learn new features on top of the input features and thus the network will have better performance (<xref ref-type="bibr" rid="B12">Chollet, 2017</xref>).</p>
<p>The BottleNeck residual module is the basic unit of ResNet (<xref ref-type="bibr" rid="B25">Han et al. 2015</xref>). It consists of three convolutional layers as shown in <xref ref-type="fig" rid="F11">Figure 11</xref>. The number of channels at input is restored by reducing the dimensionality using 1&#x2a;1 convolution, then using 3&#x2a;3 convolution for feature extraction, and finally using 1&#x2a;1 convolution to raise the dimensionality.</p>
<fig id="F11" position="float">
<label>FIGURE 11</label>
<caption>
<p>Bottleneck residual module.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g011.tif"/>
</fig>
<p>The basic module of SqueezeNext (<xref ref-type="bibr" rid="B20">Gholami et al., 2018</xref>) is adapted from BottleNeck. The 3&#x2a;3 convolution is factorized into the sum of two 2D convolution operations, 3&#x2a;1 and 1&#x2a;3, as shown in <xref ref-type="fig" rid="F12">Figure 12</xref>, and a set of 1&#x2a;1 convolutions is added before the start of the BottleNeck structure, which is used to reduce the number of channels. By adjusting the number of 1&#x2a;1 convolutions, the number of channels in each layer of convolution can be flexibly controlled to ensure that the two convolution operations obtained from the 3&#x2a;3 convolution factorization are always in the lower dimensional channels. The final 1&#x2a;1 convolution in the module is used to restore the features to the same dimension as the input channels.</p>
<fig id="F12" position="float">
<label>FIGURE 12</label>
<caption>
<p>SqueezeNext module.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g012.tif"/>
</fig>
<p>ShuffleNetV1 (<xref ref-type="bibr" rid="B97">Zhang et al., 2018</xref>) is optimized based on the bottleneck structure, as shown in <xref ref-type="fig" rid="F13">Figure 13</xref>, replacing the normal 3&#x2a;3 convolution with a 3&#x2a;3 DW convolution, replacing the 1&#x2a;1 convolution with a 1&#x2a;1 grouping convolution, and adding a channel cleaning operation after the first 1&#x2a;1 grouping convolution. Then the original element sum operation is converted to channel cascade concat. 3&#x2a;3 DW convolution can significantly reduce the number of parameters, but when the number of channels is too high using 1&#x2a;1 convolution many times will increase the amount of computation. Group convolution perfectly solves this problem, after the experimental demonstration, 0.25 times the number of groups tend to sustain better results, which indicates that wider feature maps can bring better results for smaller models.</p>
<fig id="F13" position="float">
<label>FIGURE 13</label>
<caption>
<p>ShuffleNet module.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g013.tif"/>
</fig>
<p>The spatial bottleneck module of DetNet (<xref ref-type="bibr" rid="B47">Li et al., 2018</xref>) is also an improvement of the bottleneck structure, which replaces the original 3&#x2a;3 convolution with the cavity convolution, allowing the feeling field to be expanded arbitrarily without introducing additional parameters. The spatial bottleneck module can greatly improve the ability of the model to localize segmentation, and at the same time obtain the very important multi-scale information in the vision task. For ultrasound images that are susceptible to noise, the spatial bottleneck module allows for better feature extraction.</p>
<p>MobileNetV2 <xref ref-type="bibr" rid="B69">Sandler et al. (2018)</xref> proposes a reverse residual module, as shown in <xref ref-type="fig" rid="F14">Figure 14</xref>, which first lowers the dimensions, then raises them, and replaces the ReLU activation with a linear activation. The inverse residual module has both the optimization characteristics of the bottleneck structure and disassembles the convolution through the depth separable convolution, using a lighter weight convolution to further reduce the model computation. The ghost bottleneck with a stride of two in GhostNet is a standard residual module structure. A depth-separated convolution module is added between the two ghost modules to reduce the amount of computation (<xref ref-type="bibr" rid="B24">Han et al., 2020</xref>).</p>
<fig id="F14" position="float">
<label>FIGURE 14</label>
<caption>
<p>Inverse residual module.</p>
</caption>
<graphic xlink:href="fphys-16-1536542-g014.tif"/>
</fig>
<p>
<xref ref-type="bibr" rid="B39">Khan et al. (2024)</xref> proposed ESDMR-Net, a deep convolutional network architecture with Squeeze-Excitation (SE) module, to better handle high-frequency information and feature variations in bruise images. The SE module extends the multi-scale information through deep separable convolution and extracts compact salient information through bottleneck layers to enhance the feature representation capability. The deep structure of the network enables feature refinement and accumulation through multi-branching design and reuse of SE modules, making it more robust in dealing with highly variable features.</p>
<p>Residual modules are widely used in building lightweight neural networks, often replacing standard convolutions with more lightweight depth-separable convolutions to further reduce computation on the basis of optimized structure (<xref ref-type="bibr" rid="B64">Qin et al., 2024</xref>). However, the trade-off is that the embedding of the residual module will make the original model structure relatively complex, and it is often necessary to utilize model compression to further reduce the memory footprint of the model.</p>
</sec>
<sec id="s3-2-2">
<title>3.2.2 Grouped convolution</title>
<p>Originating in 2012, AlexNet (<xref ref-type="bibr" rid="B41">Krizhevsky et al., 2012</xref>) split the convolution, and grouped convolution effectively reduced the computational complexity (FLOPs) and achieved structured sparsity. When the number of groups is equal to the number of input channels, grouped convolution can be transformed into a DW convolution to further reduce parameters. Each branch of the ResNeXt module (<xref ref-type="bibr" rid="B86">Xie et al., 2017</xref>) employed an identical convolutional topology. The idea of grouped convolution is actually to adjust the number of channels involved in the convolution operation. By splitting the number of channels for convolution, each grouping is executed in parallel. The number of subgroups depends on the hardware resource configuration currently in use and is related to the design of the network structure being embraced. The design of grouped convolution is more in line with GPU hardware design principles and thus runs faster than the Inception (<xref ref-type="bibr" rid="B26">He et al., 2016</xref>) module with manually designed convolutional details.</p>
<p>Grouped convolution allows more channels to be used with fixed FLOPs and increases the capacity of the network, so networks with grouped convolution (<xref ref-type="bibr" rid="B31">Howard, 2017</xref>; <xref ref-type="bibr" rid="B97">Zhang et al., 2018</xref>; <xref ref-type="bibr" rid="B86">Xie et al., 2017</xref>; <xref ref-type="bibr" rid="B55">Ma et al., 2018</xref>; <xref ref-type="bibr" rid="B22">Guo et al., 2022</xref>) can maintain high accuracy while reducing FLOPs. However, increasing the number of channels leads to higher memory access costs (MACs). ShuffleNetV2 (<xref ref-type="bibr" rid="B55">Ma et al., 2018</xref>) introduces a channel splitting module in its basic unit, which functions similarly to grouped convolution but helps mitigate the excessive MACs associated with an increased number of grouped convolutions.</p>
<p>The Ghost module in GhostNet (<xref ref-type="bibr" rid="B24">Han et al., 2020</xref>) divides the results generated by convolution into two groups. One group retains part of the original convolution, and the other group is optimized for the redundant features in the output that are similar to each other. A small number of base features are first generated from some of the standard convolutional layers. Secondly, a series of linear transformations are used to further generate new features, which are called Ghost feature maps. A large number of similar redundant feature maps generated by the convolution operation are replaced by Ghost features, achieving model speedup.</p>
<p>Grouped convolution allows different groups of channels to focus on different types of features, separating local features from global features. For example, certain groups focus on edge information, while others focus on texture or shape, avoiding interference from scattered noise in all channels.</p>
</sec>
<sec id="s3-2-3">
<title>3.2.3 SE module</title>
<p>SE module (Squeeze-Excitation) (<xref ref-type="bibr" rid="B32">Hu et al., 2018</xref>) includes two operations: Squeeze and Excitation. It is hoped that the model can autonomously learn the dependencies between different channels and obtain the relationships between features. Squeeze uses a 1&#x2a;1 global average pooling operation to compress the spatial features of each channel into a scalar to obtain a global description of the channel. Excitation performs dynamic adaptive adjustment of the channel weights. The global description generated in the Squeeze part is fed into a sub-network with two fully connected layers for learning the dependencies between the channels. The global description obtained in the Squeeze part amplifies the sensory range, avoiding the limitation that small sensory fields in the shallow network are unable to sense more features. The Excitation part is an automated gating mechanism that allows the model to adaptively focus on the important feature channels.</p>
<p>The SE module is portable and can provide significant performance improvements for deeper architectures with minimal additional computational cost. The SE module can be embedded in the ResNet residual network (<xref ref-type="bibr" rid="B26">He et al., 2016</xref>), where it is placed after the output of the main branch of each residual module, and the weighted outputs are obtained after Squeeze and Excitation. The main branch features processed by the SE module are then added with the residual branches to obtain the final output features. In MobileNetV3 (<xref ref-type="bibr" rid="B30">Howard et al., 2019</xref>), the SE module is added after the DW convolution block, further improving accuracy without increasing time loss.</p>
<p>The embedding of the SE module can help the model to better learn the information of each channel and enhance the network&#x2019;s ability to pay attention to the key features when recognizing the dynamically changing time-domain information of ultrasound signals and improve the robustness of the anomalous signal detection. <xref ref-type="bibr" rid="B36">Jiang et al. (2019)</xref> combined the residual module and the Squeeze-and-Excitation module to design a small SE-ResNet module for the classification of breast cancer histopathology images to reduce the training parameters of the model and the risk of over-fitting. <xref ref-type="bibr" rid="B96">Zhang et al. (2020)</xref> used a similar approach to solve the problem that the model cannot extract accurate features of long-term sequences in the task of signal classification.</p>
</sec>
<sec id="s3-2-4">
<title>3.2.4 Attention mechanisms module</title>
<p>Attention Mechanism Modules (<xref ref-type="bibr" rid="B22">Guo et al., 2022</xref>) are lightweight and generalized modules that allow feature focusing in the channel dimension and spatial dimension, and integration of independent dimensions of both. Thus,the attention mechanism module can be categorized into channel attention module (CAM) and spatial attention module (SAM). The attention mechanism module is similar to the SE module, both CAM and SAM use a double-pooling operation, that is, adding a global maximum pooling operation on top of the global average pooling. This dual-pool operation structure can extract richer high-level feature information. The encapsulated attention mechanism module can be directly embedded behind the regular convolutional layer of the feedforward neural network without any additional computational overhead.</p>
<p>The Coordinate Attention (CA) module (<xref ref-type="bibr" rid="B29">Hou et al., 2021</xref>) generates two 1D feature representations by aggregating the global information of the input feature map in the height and width directions, respectively. Unlike traditional attention mechanisms, the CA module captures the dependencies of features in both vertical and horizontal directions, making it easier for the model to capture the exact location of the target. Therefore, embedding the CA module in the model that performs ultrasound signal classification can help the model more accurately identify the location of the disease and improve the accuracy of diagnosis (<xref ref-type="bibr" rid="B3">Ansari et al., 2024</xref>).</p>
<p>MobileVit (<xref ref-type="bibr" rid="B51">Liu et al., 2017</xref>) integrates the Transformer module into a lightweight convolutional neural network, which retains the efficient local feature extraction capability of the network and enhances the capture of global relationships. By combining the advantages of both, MobileVit can be adapted to multitasking scenarios, and its efficient global feature capture capability can be used in complex scenarios such as ultrasound signal processing.</p>
<p>Denoising of ultrasound images may result in the loss of local detail features. SBCFormer, proposed by <xref ref-type="bibr" rid="B54">Lu et al. (2024)</xref>, solved this problem by designing a two-stream block structure. One stream is used to reduce the feature map size and apply the attention mechanism, while the other stream is used as a &#x201c;pass-through channel&#x201d; to retain the local information of the input feature map.</p>
</sec>
</sec>
<sec id="s3-3">
<title>3.3 Analysis and summary</title>
<p>This subsection summarises the lightweight convolutional neural networks. <xref ref-type="table" rid="T2">Table 2</xref> demonstrates the performance of the lightweight models on the ImageNet dataset. In order to fix a benchmark for comparison, these models are evaluated on the ImageNet dataset, using the parametric counts, FLOPs, and MAC metrics to measure the lightweight effect and also focusing on the classification accuracy of the models. Analogising the performance of these models on the ImageNet dataset can provide ideas for the ultrasound signal analysis task. FLOPs and a number of covariates metrics do not fully reflect the actual efficiency of the models (<xref ref-type="bibr" rid="B43">Lai et al., 2018</xref>) and still need to be optimised according to the task scenario. <xref ref-type="table" rid="T3">Table 3</xref> demonstrates the lightweight techniques used in these models, and it is easy to see that a combination of optimisation techniques is often required to achieve a model that maintains higher accuracy while improving computational efficiency.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>The comparison of lightweight models on ImageNet.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Model</th>
<th align="center">Parameter(M)</th>
<th align="center">FLOPs(M)</th>
<th align="center">MACs(G)</th>
<th align="center">TOP-1 ACC. (%)</th>
<th align="center">TOP-2 ACC. (%)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">SqueezeNet (2016)</td>
<td align="center">4.8</td>
<td align="center">0.82</td>
<td align="center">0.35</td>
<td align="center">57.5</td>
<td align="center">80.3</td>
</tr>
<tr>
<td align="left">InceptionV3 (2016)</td>
<td align="center">23.62</td>
<td align="center">
<inline-formula id="inf36">
<mml:math id="m37">
<mml:mrow>
<mml:mn>5.72</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>1024</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">5.72</td>
<td align="center">77.9</td>
<td align="center">93.7</td>
</tr>
<tr>
<td align="left">MobileNetV1 (2017)</td>
<td align="center">4.2</td>
<td align="center">575</td>
<td align="center">0.55</td>
<td align="center">70.6</td>
<td align="center">89.5</td>
</tr>
<tr>
<td align="left">Xception (2017)</td>
<td align="center">22.85</td>
<td align="center">
<inline-formula id="inf37">
<mml:math id="m38">
<mml:mrow>
<mml:mn>8.4</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>1024</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">8.42</td>
<td align="center">78.8</td>
<td align="center">94.3</td>
</tr>
<tr>
<td align="left">SqueezeNeXt (2018)</td>
<td align="center">3.2</td>
<td align="center">708</td>
<td align="center">0.69</td>
<td align="center">67.5</td>
<td align="center">88.2</td>
</tr>
<tr>
<td align="left">MobileNetV2 (2018)</td>
<td align="center">3.4</td>
<td align="center">300</td>
<td align="center">0.29</td>
<td align="center">72.0</td>
<td align="center">91.0</td>
</tr>
<tr>
<td align="left">ShuffleNet (2018)</td>
<td align="center">3.46</td>
<td align="center">140</td>
<td align="center">0.14</td>
<td align="center">72.6</td>
<td align="center">&#x2014;</td>
</tr>
<tr>
<td align="left">ShuffleNetV2 (2018)</td>
<td align="center">2.3</td>
<td align="center">146</td>
<td align="center">0.15</td>
<td align="center">71.8</td>
<td align="center">&#x2014;</td>
</tr>
<tr>
<td align="left">PeleeNet (2018)</td>
<td align="center">2.8</td>
<td align="center">508</td>
<td align="center">0.51</td>
<td align="center">72.6</td>
<td align="center">90.6</td>
</tr>
<tr>
<td align="left">MobileNetV3 (2019)</td>
<td align="center">3.2</td>
<td align="center">265</td>
<td align="center">0.21</td>
<td align="center">75.2</td>
<td align="center">&#x2014;</td>
</tr>
<tr>
<td align="left">EfficientNetV1 (2019)</td>
<td align="center">5.3</td>
<td align="center">390</td>
<td align="center">0.18</td>
<td align="center">77.1</td>
<td align="center">93.3</td>
</tr>
<tr>
<td align="left">GhostNet (2020)</td>
<td align="center">5.2</td>
<td align="center">141</td>
<td align="center">0.14</td>
<td align="center">73.9</td>
<td align="center">94.1</td>
</tr>
<tr>
<td align="left">EfficientNetV2 (2021)</td>
<td align="center">24</td>
<td align="center">280</td>
<td align="center">0.28</td>
<td align="center">83.9</td>
<td align="center">&#x2014;</td>
</tr>
<tr>
<td align="left">VanillaNet (2023)</td>
<td align="center">15.5</td>
<td align="center">520</td>
<td align="center">0.52</td>
<td align="center">72.49</td>
<td align="center">79.66</td>
</tr>
<tr>
<td align="left">MFDNet (2023)</td>
<td align="center">11.46</td>
<td align="center">384</td>
<td align="center">0.38</td>
<td align="center">77.8</td>
<td align="center">&#x2014;</td>
</tr>
<tr>
<td align="left">SBCFormer (2024)</td>
<td align="center">5.6</td>
<td align="center">700</td>
<td align="center">0.7</td>
<td align="center">75.8</td>
<td align="center">&#x2014;</td>
</tr>
<tr>
<td align="left">MobileNetV4 (2024)</td>
<td align="center">9.2</td>
<td align="center">220</td>
<td align="center">0.22</td>
<td align="center">82.9</td>
<td align="center">&#x2014;</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Models and their lightweight technologies.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Model</th>
<th align="center">LightWeight technology</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">SqueezeNet (2016)</td>
<td align="center">Calculation decoupling</td>
</tr>
<tr>
<td align="left">InceptionV3 (2016)</td>
<td align="center">Calculation decoupling</td>
</tr>
<tr>
<td align="left">MobileNetV1 (2017)</td>
<td align="center">Tensor decoupling (Depthwise separable convolution)</td>
</tr>
<tr>
<td align="left">Xception (2017)</td>
<td align="center">Tensor decoupling (Depthwise separable convolution)</td>
</tr>
<tr>
<td align="left">SqueezeNeXt (2018)</td>
<td align="center">Tensor decoupling (Depthwise separable convolution); Residual module</td>
</tr>
<tr>
<td align="left">MobileNetV2 (2018)</td>
<td align="center">Tensor decoupling (Depthwise separable convolution); Residual module (Inverted residuals)</td>
</tr>
<tr>
<td align="left">ShuffleNet (2018)</td>
<td align="center">Group convolution; Channel shuffle</td>
</tr>
<tr>
<td align="left">ShuffleNetV2 (2018)</td>
<td align="center">Group convolution; Residual module (Inverted residuals)</td>
</tr>
<tr>
<td align="left">PeleeNet (2018)</td>
<td align="center">Calculation decoupling; Tensor decoupling (Depthwise separable convolution)</td>
</tr>
<tr>
<td align="left">MobileNetV3 (2019)</td>
<td align="center">SE module; Tensor decoupling (Depthwise separable convolution)</td>
</tr>
<tr>
<td align="left">GhostNet (2020)</td>
<td align="center">Tensor decoupling (Depthwise separable convolution)</td>
</tr>
<tr>
<td align="left">EfficientNetV2 (2021)</td>
<td align="center">Tensor decoupling (Depthwise separable convolution)</td>
</tr>
<tr>
<td align="left">VanillaNet (2023)</td>
<td align="center">SE module; Attention mechanisms module</td>
</tr>
<tr>
<td align="left">MFDNet (2023)</td>
<td align="center">Residual module; Attention mechanisms module</td>
</tr>
<tr>
<td align="left">SBCFormer (2024)</td>
<td align="center">Residual module (Inverted residuals); Attention mechanisms module</td>
</tr>
<tr>
<td align="left">MobileNetV4 (2024)</td>
<td align="center">Attention mechanisms module</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
</sec>
<sec id="s4">
<title>4 Future research directions</title>
<p>Based on the above analysis, the current lightweight technology has limitations in its wide application to a certain extent. Future research should focus on the following promising directions:</p>
<sec id="s4-1">
<title>4.1 Data imbalance</title>
<p>Designing a lightweight model for ultrasound signals requires overcoming the impact of experimental data imbalance on the model (<xref ref-type="bibr" rid="B61">Mone et al., 2023</xref>). Multiple analysis methods can be used to reduce the bias that may be caused by a single experimental method. For lightweight models, methods such as adjusting the loss function and category weights can be used to give a larger weight to minority categories, thereby reducing the impact of data imbalance on model performance and improving the model&#x2019;s ability to identify different categories.</p>
</sec>
<sec id="s4-2">
<title>4.2 Lack of statistically significant clinical efficacy</title>
<p>Existing AI-assisted diagnosis and treatment methods are basically still in the research stage, lacking extensive clinical trials, and there is a gap between translational research and clinical application (<xref ref-type="bibr" rid="B15">Dhage et al., 2021</xref>). In the future, large-scale randomized controlled trials (RCTs) can be conducted to evaluate the clinical effectiveness of the model. Additionally, different organs and tissues have different ultrasonic features, such as periodic heartbeat signals, non-linear muscle tissue echoes, and time domain features depend on dynamic changes within the time window. Therefore more extensive evaluation across diverse ultrasound datasets, incorporating different probes, imaging frequencies, and tissue types, is necessary to validate the generalizability of these models. Multi-center studies and cross-dataset validation will be critical to ensuring robust performance in varied clinical settings.</p>
</sec>
<sec id="s4-3">
<title>4.3 Interpretability and generalizability of the model</title>
<p>The interpretability of model research requires sufficient theoretical guidance and experimental analysis. Many studies are based on small-scale data sets or data from specific institutions and lack multicenter, multiethnic, and multienvironmental data validation (<xref ref-type="bibr" rid="B3">Ansari et al., 2024</xref>). More ablation experiments are necessary to focus more on which part of the lightweight CNN leads to performance improvement. This is a key factor to accurately improve image recognition performance, rather than blindly stacking techniques to improve performance (<xref ref-type="bibr" rid="B7">Chandrasekar et al., 2022</xref>). Future research should focus on optimizing lightweight CNNs for real-time ultrasound applications.</p>
</sec>
</sec>
<sec sec-type="conclusion" id="s5">
<title>5 Conclusion</title>
<p>This paper presents a detailed review of the current state of research and the challenges of lightweight techniques for the task of ultrasound signal classification. Pruning and knowledge distillation techniques improve diagnostic accuracy while reducing model complexity, especially structured pruning that removes redundant filters in M-mode ultrasound and focuses on critical temporal features of the time-domain signal. Operational optimization techniques optimize computational efficiency while improving feature extraction capabilities, adapting to the deployment of embedded devices. In addition, for the scattering noise and acoustic shadow effect in ultrasound signals, adding feature enhancement modules to the network can effectively improve the robustness and classification accuracy of the model. In the future, based on the lightweight model architecture, we will use model compression techniques to design model architectures specifically for M-mode ultrasound and time series signals, which will further promote their application in diverse clinical scenarios such as fetal health detection and lung disease diagnosis. Through continuous research and innovation, lightweight technology will become an important bridge connecting AI with the practical needs of healthcare.</p>
</sec>
</body>
<back>
<sec sec-type="author-contributions" id="s6">
<title>Author contributions</title>
<p>BZ: Writing &#x2013; original draft, Writing &#x2013; review and editing, Conceptualization, Data curation, Formal Analysis, Methodology, Software, Visualization. ZL: Writing &#x2013; review and editing, Supervision. YH: Writing &#x2013; review and editing. LW: Supervision, Writing &#x2013; review and editing. XL: Writing &#x2013; review and editing. YY: Writing &#x2013; review and editing.</p>
</sec>
<sec sec-type="funding-information" id="s7">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research and/or publication of this article. This work was supported by the grograms with No. NMKY14520135 and No.2024NCUTYXCX201.</p>
</sec>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="s9">
<title>Generative AI statement</title>
<p>The author(s) declare that no Generative AI was used in the creation of this manuscript.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Afsa</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Ansari</surname>
<given-names>M. Y.</given-names>
</name>
<name>
<surname>Paul</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Halabi</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Alataresh</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Shah</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>Development and validation of a class imbalance-resilient cardiac arrest prediction framework incorporating multiscale aggregation, ica and explainability</article-title>. <source>IEEE Trans. Biomed. Eng.</source> <pub-id pub-id-type="doi">10.1109/TBME.2024.3517635</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ansari</surname>
<given-names>M. Y.</given-names>
</name>
<name>
<surname>Mangalote</surname>
<given-names>I. A. C.</given-names>
</name>
<name>
<surname>Masri</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Dakua</surname>
<given-names>S. P.</given-names>
</name>
</person-group> (<year>2023</year>). &#x201c;<article-title>Neural network-based fast liver ultrasound image segmentation</article-title>,&#x201d; in <source>2023 international joint conference on neural networks (IJCNN)</source>. <publisher-name>IEEE</publisher-name>, <fpage>1</fpage>&#x2013;<lpage>8</lpage>.</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ansari</surname>
<given-names>M. Y.</given-names>
</name>
<name>
<surname>Mangalote</surname>
<given-names>I. A. C.</given-names>
</name>
<name>
<surname>Meher</surname>
<given-names>P. K.</given-names>
</name>
<name>
<surname>Aboumarzouk</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Al-Ansari</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Halabi</surname>
<given-names>O.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>Advancements in deep learning for b-mode ultrasound segmentation: a comprehensive review</article-title>. <source>IEEE Trans. Emerg. Top. Comput. Intell.</source> <volume>8</volume>, <fpage>2126</fpage>&#x2013;<lpage>2149</lpage>. <pub-id pub-id-type="doi">10.1109/tetci.2024.3377676</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ayle</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Charpentier</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Rachwan</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Z&#xfc;gner</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Geisler</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>G&#xfc;nnemann</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>On the robustness and anomaly detection of sparse neural networks</article-title>. <source>arXiv Prepr. arXiv:2207.04227</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2207.04227</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bria</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Marrocco</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Tortorella</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Addressing class imbalance in deep learning for small lesion detection on medical images</article-title>. <source>Comput. Biol. Med.</source> <volume>120</volume>, <fpage>103735</fpage>. <pub-id pub-id-type="doi">10.1016/j.compbiomed.2020.103735</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Burch</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Idumah</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Doga</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Lartey</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Yehia</surname>
<given-names>L.</given-names>
</name>
<etal/>
</person-group> (<year>2025</year>). <article-title>Towards quantum tensor decomposition in biomedical applications</article-title>. <source>arXiv Prepr. arXiv:2502.13140</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2502.13140</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chandrasekar</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Singh</surname>
<given-names>A. V.</given-names>
</name>
<name>
<surname>Maharjan</surname>
<given-names>R. S.</given-names>
</name>
<name>
<surname>Dakua</surname>
<given-names>S. P.</given-names>
</name>
<name>
<surname>Balakrishnan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Dash</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Perspectives on the technological aspects and biomedical applications of virus-like particles/nanoparticles in reproductive biology: insights on the medicinal and toxicological outlook</article-title>. <source>Adv. NanoBiomed Res.</source> <volume>2</volume>, <fpage>2200010</fpage>. <pub-id pub-id-type="doi">10.1002/anbr.202200010</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Ren</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Review of lightweight deep convolutional neural networks</article-title>. <source>Archives Comput. Methods Eng.</source> <volume>31</volume>, <fpage>1915</fpage>&#x2013;<lpage>1937</lpage>. <pub-id pub-id-type="doi">10.1007/s11831-023-10032-z</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Tao</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Vanillanet: the power of minimalism in deep learning</article-title>. <source>Adv. Neural Inf. Process. Syst.</source> <volume>36</volume>, <fpage>7050</fpage>&#x2013;<lpage>7064</lpage>. <pub-id pub-id-type="doi">10.48550/arXiv.2305.12972</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>de Hoog</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Improved feature distillation via projector ensemble</article-title>. <source>Adv. Neural Inf. Process. Syst.</source> <volume>35</volume>, <fpage>12084</fpage>&#x2013;<lpage>12095</lpage>. <pub-id pub-id-type="doi">10.48550/arXiv.2210.15274</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>A survey of model compression and acceleration for deep neural networks</article-title>. <source>arXiv Prepr. arXiv:1710.09282</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1710.09282</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Chollet</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>Xception: deep learning with depthwise separable convolutions</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>1251</fpage>&#x2013;<lpage>1258</lpage>.</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dakua</surname>
<given-names>S. P.</given-names>
</name>
<name>
<surname>Abinahed</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Zakaria</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Balakrishnan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Younes</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Navkar</surname>
<given-names>N.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Moving object tracking in clinical scenarios: application to cardiac surgery and cerebral aneurysm clipping</article-title>. <source>Int. J. Comput. assisted radiology Surg.</source> <volume>14</volume>, <fpage>2165</fpage>&#x2013;<lpage>2176</lpage>. <pub-id pub-id-type="doi">10.1007/s11548-019-02030-z</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Deffieux</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Demen&#xe9;</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Tanter</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Functional ultrasound imaging: a new imaging modality for neuroscience</article-title>. <source>Neuroscience</source> <volume>474</volume>, <fpage>110</fpage>&#x2013;<lpage>121</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuroscience.2021.03.005</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dhage</surname>
<given-names>P. A.</given-names>
</name>
<name>
<surname>Sharbidre</surname>
<given-names>A. A.</given-names>
</name>
<name>
<surname>Dakua</surname>
<given-names>S. P.</given-names>
</name>
<name>
<surname>Balakrishnan</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Leveraging hallmark alzheimer&#x2019;s molecular targets using phytoconstituents: current perspective and emerging trends</article-title>. <source>Biomed. and Pharmacother.</source> <volume>139</volume>, <fpage>111634</fpage>. <pub-id pub-id-type="doi">10.1016/j.biopha.2021.111634</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Djelouat</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Zhai</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Al Disi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Amira</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bensaali</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>System-on-chip solution for patients biometric: a compressive sensing-based approach</article-title>. <source>IEEE Sensors J.</source> <volume>18</volume>, <fpage>9629</fpage>&#x2013;<lpage>9639</lpage>. <pub-id pub-id-type="doi">10.1109/jsen.2018.2871411</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Elhoushi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Shafiq</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>Y. H.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J. Y.</given-names>
</name>
</person-group> (<year>2021</year>). &#x201c;<article-title>Deepshift: towards multiplication-less neural networks</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</source>, <fpage>2359</fpage>&#x2013;<lpage>2368</lpage>.</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Esfahani</surname>
<given-names>S. S.</given-names>
</name>
<name>
<surname>Zhai</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Amira</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bensaali</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>AbiNahed</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Lattice-Boltzmann interactive blood flow simulation pipeline</article-title>. <source>Int. J. Comput. assisted radiology Surg.</source> <volume>15</volume>, <fpage>629</fpage>&#x2013;<lpage>639</lpage>. <pub-id pub-id-type="doi">10.1007/s11548-020-02120-3</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Filters&#xe2;&#x20ac;&#x2122;Importance</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Pruning filters for efficient convnets</article-title>
</citation>
</ref>
<ref id="B20">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Gholami</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Kwon</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Tai</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Yue</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>P.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). &#x201c;<article-title>Squeezenext: hardware-aware neural network design</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition workshops</source>, <fpage>1638</fpage>&#x2013;<lpage>1647</lpage>.</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Goo&#xdf;en</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Deshpande</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Harder</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Schwab</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Baltruschat</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Mabotuwana</surname>
<given-names>T.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). &#x201c;<article-title>Deep learning for pneumothorax detection and localization in chest radiographs</article-title>,&#x201d; <source>arXiv Prepr. arXiv:1907.07324</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1907.07324</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guo</surname>
<given-names>M. H.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>T. X.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>J. J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Z. N.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>P. T.</given-names>
</name>
<name>
<surname>Mu</surname>
<given-names>T. J.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Attention mechanisms in computer vision: a survey</article-title>. <source>Comput. Vis. media</source> <volume>8</volume>, <fpage>331</fpage>&#x2013;<lpage>368</lpage>. <pub-id pub-id-type="doi">10.1007/s41095-022-0271-y</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Gupta</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bau</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Jha</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Garud</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2024</year>). &#x201c;<article-title>Torque based structured pruning for deep neural network</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF winter conference on applications of computer vision</source>, <fpage>2711</fpage>&#x2013;<lpage>2720</lpage>.</citation>
</ref>
<ref id="B24">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Han</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>Ghostnet: more features from cheap operations</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</source>, <fpage>1580</fpage>&#x2013;<lpage>1589</lpage>.</citation>
</ref>
<ref id="B25">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Han</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Mao</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Dally</surname>
<given-names>W. J.</given-names>
</name>
</person-group> (<year>2015</year>). &#x201c;<article-title>Deep compression: compressing deep neural networks with pruning</article-title>,&#x201d; in <source>Trained quantization and huffman coding</source>. <publisher-loc>San Juan, Puerto Rico</publisher-loc>: <publisher-name>International Conference on Learning Representations (ICLR) 2016</publisher-name>. <comment>arXiv preprint arXiv:1510.00149</comment>.</citation>
</ref>
<ref id="B26">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Ren</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Deep residual learning for image recognition</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>770</fpage>&#x2013;<lpage>778</lpage>.</citation>
</ref>
<ref id="B27">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Heo</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yun</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Kwak</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Choi</surname>
<given-names>J. Y.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>A comprehensive overhaul of feature distillation</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF international conference on computer vision</source>, <fpage>1921</fpage>&#x2013;<lpage>1930</lpage>.</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hinton</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Distilling the knowledge in a neural network</article-title>. <source>arXiv Prepr. arXiv:1503.02531</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1503.02531</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Hou</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Feng</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). &#x201c;<article-title>Coordinate attention for efficient mobile network design</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</source>, <fpage>13713</fpage>&#x2013;<lpage>13722</lpage>.</citation>
</ref>
<ref id="B30">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Howard</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Sandler</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Chu</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>L. C.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Tan</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). &#x201c;<article-title>Searching for mobilenetv3</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF international conference on computer vision</source>, <fpage>1314</fpage>&#x2013;<lpage>1324</lpage>.</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Howard</surname>
<given-names>A. G. M.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Efficient convolutional neural networks for mobile vision applications</article-title>. <source>arXiv Prepr. arXiv:1704.04861</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1704.04861</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Hu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Squeeze-and-excitation networks</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>7132</fpage>&#x2013;<lpage>7141</lpage>.</citation>
</ref>
<ref id="B33">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Hua</surname>
<given-names>B. S.</given-names>
</name>
<name>
<surname>Tran</surname>
<given-names>M. K.</given-names>
</name>
<name>
<surname>Yeung</surname>
<given-names>S. K.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Pointwise convolutional neural networks</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>984</fpage>&#x2013;<lpage>993</lpage>.</citation>
</ref>
<ref id="B34">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>D. S.</given-names>
</name>
<name>
<surname>Premaratne</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Qu</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Jo</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Hussain</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2023</year>). &#x201c;<article-title>Advanced intelligent computing technology and applications</article-title>,&#x201d; in <source>19th international conference</source>. <publisher-loc>Zhengzhou, China</publisher-loc>: <publisher-name>International Conference on Intelligent Computing, ICIC 2023</publisher-name>, <fpage>10</fpage>&#x2013;<lpage>13</lpage>.</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Iandola</surname>
<given-names>F. N.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Squeezenet: alexnet-level accuracy with 50x fewer parameters and&#xa1; 0.5 mb model size</article-title>,&#x201d;. <source>arXiv Prepr. arXiv:1602.07360</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1602.07360</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jiang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Xiao</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Breast cancer histopathological image classification using convolutional neural networks with small se-resnet module</article-title>. <source>PloS one</source> <volume>14</volume>, <fpage>e0214587</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0214587</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Jin</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2023</year>). &#x201c;<article-title>Multi-level logit distillation</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</source>, <fpage>24276</fpage>&#x2013;<lpage>24285</lpage>.</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>J. Y.</given-names>
</name>
<name>
<surname>Yoo</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>A new feature-enhanced speckle reduction method based on multiscale analysis for ultrasound b-mode imaging</article-title>. <source>IEEE Trans. Biomed. Eng.</source> <volume>63</volume>, <fpage>1178</fpage>&#x2013;<lpage>1191</lpage>. <pub-id pub-id-type="doi">10.1109/TBME.2015.2486042</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Khan</surname>
<given-names>T. M.</given-names>
</name>
<name>
<surname>Naqvi</surname>
<given-names>S. S.</given-names>
</name>
<name>
<surname>Meijering</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Esdmr-net: a lightweight network with expand-squeeze and dual multiscale residual connections for medical image segmentation</article-title>. <source>Eng. Appl. Artif. Intell.</source> <volume>133</volume>, <fpage>107995</fpage>. <pub-id pub-id-type="doi">10.1016/j.engappai.2024.107995</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kolda</surname>
<given-names>T. G.</given-names>
</name>
<name>
<surname>Bader</surname>
<given-names>B. W.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Tensor decompositions and applications</article-title>. <source>SIAM Rev.</source> <volume>51</volume>, <fpage>455</fpage>&#x2013;<lpage>500</lpage>. <pub-id pub-id-type="doi">10.1137/07070111x</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Krizhevsky</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Sutskever</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Hinton</surname>
<given-names>G. E.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Imagenet classification with deep convolutional neural networks</article-title>. <source>Adv. neural Inf. Process. Syst.</source> <volume>25</volume>. <pub-id pub-id-type="doi">10.1145/3065386</pub-id>
</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kuznik</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Obrecht</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Rusaouen</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Roux</surname>
<given-names>J. J.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Lbm based flow simulation using gpu computing processor</article-title>. <source>Comput. and Math. Appl.</source> <volume>59</volume>, <fpage>2380</fpage>&#x2013;<lpage>2392</lpage>. <pub-id pub-id-type="doi">10.1016/j.camwa.2009.08.052</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lai</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Suda</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Chandra</surname>
<given-names>V.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Not all ops are created equal</article-title>. <source>arXiv preprint arXiv:1801.04326</source>
</citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Le-Khac</surname>
<given-names>P. H.</given-names>
</name>
<name>
<surname>Healy</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Smeaton</surname>
<given-names>A. F.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Contrastive representation learning: a framework and review</article-title>. <source>Ieee Access</source> <volume>8</volume>, <fpage>193907</fpage>&#x2013;<lpage>193934</lpage>. <pub-id pub-id-type="doi">10.1109/access.2020.3031549</pub-id>
</citation>
</ref>
<ref id="B45">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2022</year>). <source>Design of ultrasound-based diagnostic algorithms for pneumothorax</source>. <publisher-name>North China University of Technology</publisher-name>. <comment>Master&#x2019;s thesis</comment>.</citation>
</ref>
<ref id="B46">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Gu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Qiao</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). &#x201c;<article-title>Blueprint separable residual network for efficient image super-resolution</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</source>, <fpage>833</fpage>&#x2013;<lpage>843</lpage>.</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Deng</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Detnet: a backbone network for object detection</article-title>. <source>arXiv Prepr. arXiv:1804.06215</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1804.06215</pub-id>
</citation>
</ref>
<ref id="B48">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Liao</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Qu&#xe9;tu</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>V. T.</given-names>
</name>
<name>
<surname>Tartaglione</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2023</year>). &#x201c;<article-title>Can unstructured pruning reduce the depth in deep neural networks?</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF international conference on computer vision</source>, <fpage>1402</fpage>&#x2013;<lpage>1406</lpage>.</citation>
</ref>
<ref id="B49">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Ning</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Dynamic feature distillation and pyramid split large kernel attention network for lightweight image super-resolution</article-title>. <source>Multimedia Tools Appl.</source> <volume>83</volume>, <fpage>79963</fpage>&#x2013;<lpage>79984</lpage>. <pub-id pub-id-type="doi">10.1007/s11042-024-18501-8</pub-id>
</citation>
</ref>
<ref id="B50">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Xiong</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2023</year>). &#x201c;<article-title>Lightweight network towards real-time image denoising on mobile devices</article-title>,&#x201d; in <source>2023 IEEE international conference on image processing (ICIP) (IEEE)</source>, <fpage>2270</fpage>&#x2013;<lpage>2274</lpage>.</citation>
</ref>
<ref id="B51">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Yan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Learning efficient convolutional networks through network slimming</article-title>. <source>Proc. IEEE Int. Conf. Comput. Vis.</source>, <fpage>2736</fpage>&#x2013;<lpage>2744</lpage>. <pub-id pub-id-type="doi">10.1109/iccv.2017.298</pub-id>
</citation>
</ref>
<ref id="B52">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Wei</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Z.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). &#x201c;<article-title>Swin transformer: hierarchical vision transformer using shifted windows</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF international conference on computer vision</source>, <fpage>10012</fpage>&#x2013;<lpage>10022</lpage>.</citation>
</ref>
<ref id="B53">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lu</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Low-rank approximation, adaptation, and other tales</article-title>. <source>arXiv Prepr. arXiv:2408</source>, <fpage>05883</fpage>. <pub-id pub-id-type="doi">10.48550/arXiv.2408.05883</pub-id>
</citation>
</ref>
<ref id="B54">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Lu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Suganuma</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Okatani</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2024</year>). &#x201c;<article-title>Sbcformer: lightweight network capable of full-size imagenet classification at 1 fps on single board computers</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF winter conference on applications of computer vision</source>, <fpage>1123</fpage>&#x2013;<lpage>1133</lpage>.</citation>
</ref>
<ref id="B55">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ma</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Zheng</surname>
<given-names>H. T.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Shufflenet v2: practical guidelines for efficient cnn architecture design</article-title>,&#x201d; in <source>Proceedings of the European conference on computer vision (ECCV)</source>, <fpage>116</fpage>&#x2013;<lpage>131</lpage>.</citation>
</ref>
<ref id="B56">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Matsuyama</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Koizumi</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Nishiyama</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tsumura</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Tsukihara</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Numata</surname>
<given-names>K.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>An avoiding overlap method between acoustic shadow and organ for automated ultrasound diagnosis and treatment</article-title>,&#x201d; in <source>2022 IEEE 11th global conference on consumer electronics (GCCE)</source>. <publisher-name>IEEE</publisher-name>, <fpage>746</fpage>&#x2013;<lpage>747</lpage>.</citation>
</ref>
<ref id="B57">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mazzeo</surname>
<given-names>M. D.</given-names>
</name>
<name>
<surname>Coveney</surname>
<given-names>P. V.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Hemelb: a high performance parallel lattice-Boltzmann code for large scale fluid flow in complex geometries</article-title>. <source>Comput. Phys. Commun.</source> <volume>178</volume>, <fpage>894</fpage>&#x2013;<lpage>914</lpage>. <pub-id pub-id-type="doi">10.1016/j.cpc.2008.02.013</pub-id>
</citation>
</ref>
<ref id="B58">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Micucci</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Iula</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Recent advances in machine learning applied to ultrasound imaging</article-title>. <source>Electronics</source> <volume>11</volume>, <fpage>1800</fpage>. <pub-id pub-id-type="doi">10.3390/electronics11111800</pub-id>
</citation>
</ref>
<ref id="B59">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mischi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Bell</surname>
<given-names>M. A. L.</given-names>
</name>
<name>
<surname>Van Sloun</surname>
<given-names>R. J.</given-names>
</name>
<name>
<surname>Eldar</surname>
<given-names>Y. C.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Deep learning in medical ultrasoundfrom image formation to image analysis</article-title>. <source>IEEE Trans. Ultrasonics, Ferroelectr. Freq. Control</source> <volume>67</volume>, <fpage>2477</fpage>&#x2013;<lpage>2480</lpage>. <pub-id pub-id-type="doi">10.1109/TUFFC.2020.3026598</pub-id>
</citation>
</ref>
<ref id="B60">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mishra</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Marr</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Apprentice: using knowledge distillation techniques to improve low-precision network accuracy</article-title>. <source>arXiv Prepr. arXiv:1711.05852</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1711.05852</pub-id>
</citation>
</ref>
<ref id="B61">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mone</surname>
<given-names>N. S.</given-names>
</name>
<name>
<surname>Syed</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ravichandiran</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Kamble</surname>
<given-names>E. E.</given-names>
</name>
<name>
<surname>Pardesi</surname>
<given-names>K. R.</given-names>
</name>
<name>
<surname>Salunke-Gawali</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>Synergistic and additive effects of menadione in combination with antibiotics on multidrug-resistant staphylococcus aureus: insights from structure-function analysis of naphthoquinones</article-title>. <source>ChemMedChem</source> <volume>18</volume>, <fpage>e202300328</fpage>. <pub-id pub-id-type="doi">10.1002/cmdc.202300328</pub-id>
</citation>
</ref>
<ref id="B62">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Papyan</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Donoho</surname>
<given-names>D. L.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Prevalence of neural collapse during the terminal phase of deep learning training</article-title>. <source>Proc. Natl. Acad. Sci.</source> <volume>117</volume>, <fpage>24652</fpage>&#x2013;<lpage>24663</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.2015509117</pub-id>
</citation>
</ref>
<ref id="B63">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Park</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>No</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>Prune your model before distill it</article-title>,&#x201d; in <source>European conference on computer vision</source>. <publisher-name>Springer</publisher-name>, <fpage>120</fpage>&#x2013;<lpage>136</lpage>.</citation>
</ref>
<ref id="B64">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qin</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Leichner</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Delakis</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Fornoni</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Luo</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>F.</given-names>
</name>
<etal/>
</person-group> <article-title>Mobilenetv4-universal models for the mobile ecosystem</article-title>.<fpage>10518</fpage> (<year>2024</year>).</citation>
</ref>
<ref id="B65">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Regaya</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Amira</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Dakua</surname>
<given-names>S. P.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Development of a cerebral aneurysm segmentation method to prevent sentinel hemorrhage</article-title>. <source>Netw. Model. Analysis Health Inf. Bioinforma.</source> <volume>12</volume>, <fpage>18</fpage>. <pub-id pub-id-type="doi">10.1007/s13721-023-00412-7</pub-id>
</citation>
</ref>
<ref id="B66">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ren</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Tong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>3d ultrasonic brain imaging with deep learning based on fully convolutional networks</article-title>. <source>Sensors</source> <volume>23</volume>, <fpage>8341</fpage>. <pub-id pub-id-type="doi">10.3390/s23198341</pub-id>
</citation>
</ref>
<ref id="B67">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Reza</surname>
<given-names>A. M.</given-names>
</name>
</person-group> (<year>2004</year>). <article-title>Realization of the contrast limited adaptive histogram equalization (clahe) for real-time image enhancement</article-title>. <source>J. VLSI signal Process. Syst. signal, image video Technol.</source> <volume>38</volume>, <fpage>35</fpage>&#x2013;<lpage>44</lpage>. <pub-id pub-id-type="doi">10.1023/b:vlsi.0000028532.53893.82</pub-id>
</citation>
</ref>
<ref id="B68">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Romero</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Ballas</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Kahou</surname>
<given-names>S. E.</given-names>
</name>
<name>
<surname>Chassang</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Gatta</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Bengio</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Fitnets: hints for thin deep nets</article-title>. <source>arXiv preprint arXiv:1412.6550</source>
</citation>
</ref>
<ref id="B69">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sandler</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Howard</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Zhmoginov</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>L. C.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Mobilenetv2: inverted residuals and linear bottlenecks</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>4510</fpage>&#x2013;<lpage>4520</lpage>.</citation>
</ref>
<ref id="B70">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sashidhar</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Kwok</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Coult</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Blackwood</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Kudenchuk</surname>
<given-names>P. J.</given-names>
</name>
<name>
<surname>Bhandari</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Machine learning and feature engineering for predicting pulse presence during chest compressions</article-title>. <source>R. Soc. Open Sci.</source> <volume>8</volume>, <fpage>210566</fpage>. <pub-id pub-id-type="doi">10.1098/rsos.210566</pub-id>
</citation>
</ref>
<ref id="B71">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Shao</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Shin</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>Structured pruning for deep convolutional neural networks via adaptive sparsity regularization</article-title>,&#x201d; in <source>2022 IEEE 46th annual computers, software, and applications conference (COMPSAC)</source>. <publisher-name>IEEE</publisher-name>, <fpage>982</fpage>&#x2013;<lpage>987</lpage>.</citation>
</ref>
<ref id="B72">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Simonyan</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Zisserman</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Very deep convolutional networks for large-scale image recognition</article-title>. <source>arXiv preprint arXiv:1409.1556</source>
</citation>
</ref>
<ref id="B73">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Srinivas</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Subramanya</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Venkatesh Babu</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>Training sparse neural networks</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition workshops</source>, <fpage>138</fpage>&#x2013;<lpage>145</lpage>.</citation>
</ref>
<ref id="B74">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sun</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Lyu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Knowledge distillation with refined logits</article-title>. <source>arXiv Prepr. arXiv:2408.07703</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2408.07703</pub-id>
</citation>
</ref>
<ref id="B75">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Szegedy</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Jia</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Sermanet</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Reed</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Anguelov</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2015</year>). &#x201c;<article-title>Going deeper with convolutions</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>1</fpage>&#x2013;<lpage>9</lpage>.</citation>
</ref>
<ref id="B76">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Szegedy</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Vanhoucke</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Ioffe</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Shlens</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Wojna</surname>
<given-names>Z.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>Rethinking the inception architecture for computer vision</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>2818</fpage>&#x2013;<lpage>2826</lpage>.</citation>
</ref>
<ref id="B77">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Tan</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Le</surname>
<given-names>Q.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>Efficientnet: rethinking model scaling for convolutional neural networks</article-title>,&#x201d; in <source>International conference on machine learning</source>. <publisher-name>International Conference on Machine Learning, PMLR 2019</publisher-name>, <fpage>6105</fpage>&#x2013;<lpage>6114</lpage>.</citation>
</ref>
<ref id="B78">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Tan</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Le</surname>
<given-names>Q.</given-names>
</name>
</person-group> (<year>2021</year>). &#x201c;<article-title>Efficientnetv2: smaller models and faster training</article-title>,&#x201d; in <source>International conference on machine learning</source>. <publisher-name>International Conference on Machine Learning, PMLR 2021</publisher-name>, <fpage>10096</fpage>&#x2013;<lpage>10106</lpage>.</citation>
</ref>
<ref id="B79">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tian</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Krishnan</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Isola</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Contrastive representation distillation</article-title>. <source>arXiv Prepr. arXiv:1910</source>, <fpage>10699</fpage>. <pub-id pub-id-type="doi">10.48550/arXiv.1910.10699</pub-id>
</citation>
</ref>
<ref id="B80">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Van Sloun</surname>
<given-names>R. J.</given-names>
</name>
<name>
<surname>Cohen</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Eldar</surname>
<given-names>Y. C.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Deep learning in ultrasound imaging</article-title>. <source>Proc. IEEE</source> <volume>108</volume>, <fpage>11</fpage>&#x2013;<lpage>29</lpage>. <pub-id pub-id-type="doi">10.1109/jproc.2019.2932116</pub-id>
</citation>
</ref>
<ref id="B81">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>R. J.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Ling</surname>
<given-names>C. X.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>K.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Pelee: a real-time object detection system on mobile devices</article-title>. <source>Adv. neural Inf. Process. Syst.</source> <volume>31</volume>, <fpage>6669</fpage>&#x2013;<lpage>6678</lpage>. <pub-id pub-id-type="doi">10.48550/arXiv.1804.06882</pub-id>
</citation>
</ref>
<ref id="B82">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2021</year>). &#x201c;<article-title>Convolutional neural network pruning with structural redundancy reduction</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF conference on computer vision and pattern recognition</source>, <fpage>14913</fpage>&#x2013;<lpage>14922</lpage>.</citation>
</ref>
<ref id="B83">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Weber</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Dexl</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>R&#xfc;gamer</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Ingrisch</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2025</year>). <article-title>Posttraining network compression for 3d medical image segmentation: reducing computational efforts via tucker decomposition</article-title>. <source>Radiol. Artif. Intell.</source> <volume>7</volume>, <fpage>e240353</fpage>. <pub-id pub-id-type="doi">10.1148/ryai.240353</pub-id>
</citation>
</ref>
<ref id="B84">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wen</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Learning structured sparsity in deep neural networks</article-title>. <source>Adv. neural Inf. Process. Syst.</source> <volume>29</volume>. <pub-id pub-id-type="doi">10.48550/arXiv.1608.03665</pub-id>
</citation>
</ref>
<ref id="B85">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Wen</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Liang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). &#x201c;<article-title>Multiple low-ranks plus sparsity based tensor reconstruction for dynamic mri</article-title>,&#x201d; in <source>2018 IEEE 23rd international conference on digital signal processing (DSP) (IEEE)</source>, <fpage>1</fpage>&#x2013;<lpage>5</lpage>.</citation>
</ref>
<ref id="B86">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Xie</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Girshick</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Doll&#xe1;r</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Tu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>K.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>Aggregated residual transformations for deep neural networks</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>1492</fpage>&#x2013;<lpage>1500</lpage>.</citation>
</ref>
<ref id="B87">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yadav</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Monath</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Angell</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Zaheer</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>McCallum</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Efficient nearest neighbor search for cross-encoder models using matrix factorization</article-title>. <source>arXiv Prepr. arXiv:2210.12579</source>, <fpage>2171</fpage>&#x2013;<lpage>2194</lpage>. <pub-id pub-id-type="doi">10.18653/v1/2022.emnlp-main.140</pub-id>
</citation>
</ref>
<ref id="B88">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ye</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Chu</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Lew</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Howard</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Robust training of neural networks at arbitrary precision and sparsity</article-title>. <source>arXiv Prepr. arXiv:2409.09245</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2409.09245</pub-id>
</citation>
</ref>
<ref id="B89">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Yim</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Joo</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Bae</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>A gift from knowledge distillation: fast optimization, network minimization and transfer learning</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>4133</fpage>&#x2013;<lpage>4141</lpage>.</citation>
</ref>
<ref id="B90">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yin</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Phan</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Zang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Liao</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Yuan</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Batude: budget-aware neural network compression based on tucker decomposition</article-title>. <source>Proc. AAAI Conf. Artif. Intell.</source> <volume>36</volume>, <fpage>8874</fpage>&#x2013;<lpage>8882</lpage>. <pub-id pub-id-type="doi">10.1609/aaai.v36i8.20869</pub-id>
</citation>
</ref>
<ref id="B91">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>You</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Crebbin</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Robust adaptive estimator for filtering noise in images</article-title>. <source>IEEE Trans. Image Process.</source> <volume>4</volume>, <fpage>693</fpage>&#x2013;<lpage>699</lpage>. <pub-id pub-id-type="doi">10.1109/83.382505</pub-id>
</citation>
</ref>
<ref id="B92">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Yu</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>C. F.</given-names>
</name>
<name>
<surname>Lai</surname>
<given-names>J. H.</given-names>
</name>
<name>
<surname>Morariu</surname>
<given-names>V. I.</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>X.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). &#x201c;<article-title>Nisp: pruning networks using neuron importance score propagation</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>9194</fpage>&#x2013;<lpage>9203</lpage>.</citation>
</ref>
<ref id="B93">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zagoruyko</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Komodakis</surname>
<given-names>N.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Paying more attention to attention: improving the performance of convolutional neural networks via attention transfer</article-title>. <source>arXiv Prepr. arXiv:1612.03928</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1612.03928</pub-id>
</citation>
</ref>
<ref id="B94">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhai</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Amira</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bensaali</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Al-Shibani</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Al-Nassr</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>El-Sayed</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2019a</year>). <article-title>Zynq soc based acceleration of the lattice Boltzmann method</article-title>. <source>Concurrency Comput. Pract. Exp.</source> <volume>31</volume>, <fpage>e5184</fpage>. <pub-id pub-id-type="doi">10.1002/cpe.5184</pub-id>
</citation>
</ref>
<ref id="B95">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhai</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Esfahani</surname>
<given-names>S. S.</given-names>
</name>
<name>
<surname>Amira</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bensaali</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Abinahed</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2019b</year>). <article-title>Heterogeneous system-on-chip-based lattice-Boltzmann visual simulation system</article-title>. <source>IEEE Syst. J.</source> <volume>14</volume>, <fpage>1592</fpage>&#x2013;<lpage>1601</lpage>. <pub-id pub-id-type="doi">10.1109/jsyst.2019.2952459</pub-id>
</citation>
</ref>
<ref id="B96">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>Se-ecgnet: a multi-scale deep residual network with squeeze-and-excitation module for ecg signal classification</article-title>,&#x201d; in <source>2020 IEEE international conference on bioinformatics and biomedicine (BIBM) (IEEE)</source>, <fpage>2685</fpage>&#x2013;<lpage>2691</lpage>.</citation>
</ref>
<ref id="B97">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2018</year>). &#x201c;<article-title>Shufflenet: an extremely efficient convolutional neural network for mobile devices</article-title>,&#x201d; in <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source>, <fpage>6848</fpage>&#x2013;<lpage>6856</lpage>.</citation>
</ref>
<ref id="B98">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zhao</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Cui</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Song</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Qiu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Liang</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>Decoupled knowledge distillation</article-title>,&#x201d; in <source>Proceedings of the IEEE/CVF Conference on computer vision and pattern recognition</source>, <fpage>11953</fpage>&#x2013;<lpage>11962</lpage>.</citation>
</ref>
<ref id="B99">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Long</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Tensor rank learning in cp decomposition via convolutional neural network</article-title>. <source>Signal Process. Image Commun.</source> <volume>73</volume>, <fpage>12</fpage>&#x2013;<lpage>21</lpage>. <pub-id pub-id-type="doi">10.1016/j.image.2018.03.017</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>