<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="review-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Neurosci.</journal-id>
<journal-title>Frontiers in Neuroscience</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Neurosci.</abbrev-journal-title>
<issn pub-type="epub">1662-453X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fnins.2024.1383844</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Neuroscience</subject>
<subj-group>
<subject>Review</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Direct training high-performance deep spiking neural networks: a review of theories and methods</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Zhou</surname> <given-names>Chenlin</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2620015/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Zhang</surname> <given-names>Han</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2598210/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Yu</surname> <given-names>Liutao</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2630622/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Ye</surname> <given-names>Yumin</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2652218/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Zhou</surname> <given-names>Zhaokun</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2651837/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Huang</surname> <given-names>Liwei</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2630614/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Ma</surname> <given-names>Zhengyu</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x0002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2374106/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Fan</surname> <given-names>Xiaopeng</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2630720/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Zhou</surname> <given-names>Huihui</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2374113/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Tian</surname> <given-names>Yonghong</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1450254/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Peng Cheng Laboratory</institution>, <addr-line>Shenzhen</addr-line>, <country>China</country></aff>
<aff id="aff2"><sup>2</sup><institution>Faculty of Computing, Harbin Institute of Technology</institution>, <addr-line>Harbin</addr-line>, <country>China</country></aff>
<aff id="aff3"><sup>3</sup><institution>School of Electronic and Computer Engineering, Shenzhen Graduate School, Peking University</institution>, <addr-line>Shenzhen</addr-line>, <country>China</country></aff>
<aff id="aff4"><sup>4</sup><institution>National Key Laboratory for Multimedia Information Processing, School of Computer Science, Peking University</institution>, <addr-line>Beijing</addr-line>, <country>China</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: Priyadarshini Panda, Yale University, United States</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Lei Deng, Tsinghua University, China</p>
<p>Chenglong Zou, Peking University, China</p></fn>
<corresp id="c001">&#x0002A;Correspondence: Zhengyu Ma <email>mazhy&#x00040;pcl.ac.cn</email></corresp>
<fn fn-type="equal" id="fn001"><p>&#x02020;These authors have contributed equally to this work and share first authorship</p></fn></author-notes>
<pub-date pub-type="epub">
<day>31</day>
<month>07</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>18</volume>
<elocation-id>1383844</elocation-id>
<history>
<date date-type="received">
<day>08</day>
<month>02</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>03</day>
<month>07</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x000A9; 2024 Zhou, Zhang, Yu, Ye, Zhou, Huang, Ma, Fan, Zhou and Tian.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Zhou, Zhang, Yu, Ye, Zhou, Huang, Ma, Fan, Zhou and Tian</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<p>Spiking neural networks (SNNs) offer a promising energy-efficient alternative to artificial neural networks (ANNs), in virtue of their high biological plausibility, rich spatial-temporal dynamics, and event-driven computation. The direct training algorithms based on the surrogate gradient method provide sufficient flexibility to design novel SNN architectures and explore the spatial-temporal dynamics of SNNs. According to previous studies, the performance of models is highly dependent on their sizes. Recently, direct training deep SNNs have achieved great progress on both neuromorphic datasets and large-scale static datasets. Notably, transformer-based SNNs show comparable performance with their ANN counterparts. In this paper, we provide a new perspective to summarize the theories and methods for training deep SNNs with high performance in a systematic and comprehensive way, including theory fundamentals, spiking neuron models, advanced SNN models and residual architectures, software frameworks and neuromorphic hardware, applications, and future trends.</p></abstract>
<kwd-group>
<kwd>deep spiking neural network</kwd>
<kwd>direct training</kwd>
<kwd>transformer-based SNNs</kwd>
<kwd>residual connection</kwd>
<kwd>energy efficiency</kwd>
<kwd>high performance</kwd>
</kwd-group>
<counts>
<fig-count count="4"/>
<table-count count="4"/>
<equation-count count="21"/>
<ref-count count="181"/>
<page-count count="19"/>
<word-count count="14849"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Neuromorphic Engineering</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<title>1 Introduction</title>
<p>Regarded as the third generation of neural network (Maass, <xref ref-type="bibr" rid="B88">1997</xref>), the brain-inspired spiking neural networks (SNNs) are potential competitors to traditional artificial neural networks (ANNs) in virtue of their high biological plausibility, and low power consumption when implemented on neuromorphic hardware (Roy et al., <xref ref-type="bibr" rid="B111">2019</xref>). In particular, the utilization of binary spikes allows SNNs to adopt low-power accumulation (AC) instead of the traditional high-power multiply-accumulation (MAC), leading to significantly enhanced energy efficiency and making SNNs increasingly popular (Chen et al., <xref ref-type="bibr" rid="B20">2023</xref>).</p>
<p>There are two mainstream pathways to obtain deep SNNs: ANN-to-SNN conversion and direct training through the surrogate gradient method. Firstly, in ANN-to-SNN conversion (Cao et al., <xref ref-type="bibr" rid="B15">2015</xref>; Hunsberger and Eliasmith, <xref ref-type="bibr" rid="B54">2015</xref>; Rueckauer et al., <xref ref-type="bibr" rid="B112">2017</xref>; Bu et al., <xref ref-type="bibr" rid="B13">2022</xref>; Meng et al., <xref ref-type="bibr" rid="B89">2022</xref>; Wang Y. et al., <xref ref-type="bibr" rid="B135">2022</xref>), a pre-trained ANN is converted to an SNN by replacing the ReLU activation layers with spiking neurons and adding scaling operations like weight normalization and threshold balancing. This conversion process suffers from long converting time steps, which causes high computational consumption in practice. In addition, the converted SNNs obtained in this way are constrained by the original ANNs&#x00027; architecture and are hard to adapt to dynamic signal (DVS, DAVIS, ATIS data) processing. Thus, the direct exploration of the virtues of SNNs is limited in ANN-to-SNN conversion. Secondly, in the field of direct training, SNNs are unfolded over simulation time steps and trained with backpropagation through time (Lee et al., <xref ref-type="bibr" rid="B67">2016</xref>; Shrestha and Orchard, <xref ref-type="bibr" rid="B119">2018</xref>). Due to the non-differentiability of spiking neurons, the surrogate gradient method is employed for backpropagation (Neftci et al., <xref ref-type="bibr" rid="B94">2019</xref>; Lee et al., <xref ref-type="bibr" rid="B65">2020b</xref>; Fang et al., <xref ref-type="bibr" rid="B35">2021a</xref>,<xref ref-type="bibr" rid="B33">b</xref>; Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>). On one hand, this direct training method can handle temporal data and also achieve decent performance on large-scale static datasets, with only a few time steps. On the other hand, it can provide sufficient flexibility for designing novel architectures specifically for SNNs and exploring the properties of SNNs directly. Therefore, the direct training method has received more attention recently.</p>
<p>Given the significant benefits and rapid advancement of directly trained deep SNNs, particularly the emergence of high-performance transformer-based SNNs, this review systematically and comprehensively summarizes the theories and methods for directly trained deep SNNs. Combining theory fundamentals, spiking neuron models, advanced SNN models and residual architectures, software frameworks and neuromorphic hardware, applications, and future trends, this article offers fresh perspectives into the field of SNNs. This review is structured as follows: Section 2 presents the evolution and recent advancements in spiking neuron models. Section 3 introduces the fundamental principles of spiking neural networks. Section 4 focuses on the most recent advanced SNN models and architectures, especially transformer-based SNNs. Section 5 concludes the software frameworks for training SNNs and the development of neuromorphic hardware. Section 6 summarizes the applications of deep SNNs. Finally, Section 7 points out future research trends and concludes this review.</p>
</sec>
<sec id="s2">
<title>2 Spiking neuron models</title>
<p>LIF (Leaky Integrate-and-Fire) neuron is one of the most commonly used neurons in SNNs (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>,<xref ref-type="bibr" rid="B171">b</xref>; Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>), which is simple but retains biological characteristics (<xref ref-type="fig" rid="F1">Figure 1A</xref>). The dynamics of LIF are described as <xref ref-type="disp-formula" rid="E1">Equations (1</xref>&#x02013;<xref ref-type="disp-formula" rid="E3">3)</xref>:</p>
<disp-formula id="E1"><label>(1)</label><mml:math id="M1"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>H</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mi>V</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi><mml:mo>-</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>&#x0002B;</mml:mo><mml:mfrac><mml:mrow><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>&#x003C4;</mml:mi></mml:mrow></mml:mfrac><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>X</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>-</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>V</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi><mml:mo>-</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>-</mml:mo><mml:msub><mml:mrow><mml:mi>V</mml:mi></mml:mrow><mml:mrow><mml:mi>r</mml:mi><mml:mi>e</mml:mi><mml:mi>s</mml:mi><mml:mi>e</mml:mi><mml:mi>t</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E2"><label>(2)</label><mml:math id="M2"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>S</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mo>&#x00398;</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>H</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>-</mml:mo><mml:msub><mml:mrow><mml:mi>V</mml:mi></mml:mrow><mml:mrow><mml:mi>t</mml:mi><mml:mi>h</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E3"><label>(3)</label><mml:math id="M3"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>V</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mi>H</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mn>1</mml:mn><mml:mo>-</mml:mo><mml:mi>S</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>&#x0002B;</mml:mo><mml:msub><mml:mrow><mml:mi>V</mml:mi></mml:mrow><mml:mrow><mml:mi>r</mml:mi><mml:mi>e</mml:mi><mml:mi>s</mml:mi><mml:mi>e</mml:mi><mml:mi>t</mml:mi></mml:mrow></mml:msub><mml:mi>S</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>where &#x003C4; in <xref ref-type="disp-formula" rid="E1">Equation (1)</xref> is the membrane time constant, <italic>X</italic>[<italic>t</italic>] is the input current at time step <italic>t</italic>. <italic>V</italic><sub><italic>reset</italic></sub> represents the reset potential, <italic>V</italic><sub><italic>th</italic></sub> represents the spike firing threshold, <italic>H</italic>[<italic>t</italic>] and <italic>V</italic>[<italic>t</italic>] represent the membrane potential before and after spike firing at time step <italic>t</italic>, respectively. &#x00398;(<italic>v</italic>) is the Heaviside step function, if <italic>v</italic> &#x02265; 0 then &#x00398;(<italic>v</italic>) &#x0003D; 1, meaning a spike is generated; otherwise &#x00398;(<italic>v</italic>) &#x0003D; 0. <italic>S</italic>[<italic>t</italic>] represents whether a neuron fires a spike at time step <italic>t</italic>.</p>
<fig id="F1" position="float">
<label>Figure 1</label>
<caption><p><bold>(A)</bold> The scheme of a spiking neuron, of which the input and output are both binary spikes. <bold>(B)</bold> The sigmoid function approximates the Heaviside activation function of a spiking neuron, and its derivative can be utilized to calculate gradients during backpropagation.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnins-18-1383844-g0001.tif"/>
</fig>
<p>LIF also comes with notable limitations in practical applications. For instance, LIF needs to manually adjust the hyperparameters, such as membrane time constant &#x003C4; and firing threshold <italic>V</italic><sub><italic>th</italic></sub>, which constrains its expressiveness. In addition, LIF is simple in modeling, which limits the range of neuronal dynamics. Overall, there is a lack of diversity and flexibility in LIF, which calls for more advanced neuron models to enhance SNNs&#x00027; performance and broaden their applications. <xref ref-type="table" rid="T1">Table 1</xref> lists some recently developed spiking neuron models and their performance on typical tasks.</p>
<table-wrap position="float" id="T1">
<label>Table 1</label>
<caption><p>Overview of spiking neurons for direct training and their performance.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>Method</bold></th>
<th valign="top" align="left"><bold>Architecture</bold></th>
<th valign="top" align="left"><bold>Dataset</bold></th>
<th valign="top" align="center"><bold>Acc (%)</bold></th>
<th valign="top" align="left"><bold>Training</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">PLIF (Fang et al., <xref ref-type="bibr" rid="B33">2021b</xref>)</td>
<td valign="top" align="left">PLIF-Net</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">93.50</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">LTMD (Wang S. et al., <xref ref-type="bibr" rid="B129">2022</xref>)</td>
<td valign="top" align="left">DenseNet</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">94.19</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">GLIF (Yao et al., <xref ref-type="bibr" rid="B151">2022</xref>)</td>
<td valign="top" align="left">ResNet-34</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">95.03</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">MLF (Feng L. et al., <xref ref-type="bibr" rid="B36">2022</xref>)</td>
<td valign="top" align="left">DS ResNet</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">94.25</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">LIFB (Shen et al., <xref ref-type="bibr" rid="B116">2023</xref>)</td>
<td valign="top" align="left">ResNet-19</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">96.32</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">Deit-SNN (Rathi and Roy, <xref ref-type="bibr" rid="B107">2023</xref>)</td>
<td valign="top" align="left">VGG16</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">93.44</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">KLIF (Jiang and Zhang, <xref ref-type="bibr" rid="B55">2023</xref>)</td>
<td valign="top" align="left">CNN</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">92.52</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">MT-SNN (Wang X. et al., <xref ref-type="bibr" rid="B132">2023</xref>)</td>
<td valign="top" align="left">MT-VGG9</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">94.74</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">PSN (Fang et al., <xref ref-type="bibr" rid="B34">2023b</xref>)</td>
<td valign="top" align="left">PLIF-Net</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">95.32</td>
<td valign="top" align="left">Parallel</td>
</tr> <tr>
<td valign="top" align="left">GLIF (Yao et al., <xref ref-type="bibr" rid="B151">2022</xref>)</td>
<td valign="top" align="left">ResNet-34</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">77.35</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">Deit-SNN (Rathi and Roy, <xref ref-type="bibr" rid="B107">2023</xref>)</td>
<td valign="top" align="left">VGG16</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">69.67</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">LIFB (Shen et al., <xref ref-type="bibr" rid="B116">2023</xref>)</td>
<td valign="top" align="left">ResNet-19</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">78.31</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">MT-SNN (Wang X. et al., <xref ref-type="bibr" rid="B132">2023</xref>)</td>
<td valign="top" align="left">MT-VGG9</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">75.53</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">GLIF (Yao et al., <xref ref-type="bibr" rid="B151">2022</xref>)</td>
<td valign="top" align="left">ResNet-34</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">69.09</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">Deit-SNN (Rathi and Roy, <xref ref-type="bibr" rid="B107">2023</xref>)</td>
<td valign="top" align="left">VGG16</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">69.00</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">LIFB (Shen et al., <xref ref-type="bibr" rid="B116">2023</xref>)</td>
<td valign="top" align="left">SEW ResNet-34</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">70.02</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">PSN (Fang et al., <xref ref-type="bibr" rid="B34">2023b</xref>)</td>
<td valign="top" align="left">SEW ResNet-34</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">70.54</td>
<td valign="top" align="left">Parallel</td>
</tr> <tr>
<td valign="top" align="left">PLIF (Fang et al., <xref ref-type="bibr" rid="B33">2021b</xref>)</td>
<td valign="top" align="left">PLIF-Net</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">74.80</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">GLIF (Yao et al., <xref ref-type="bibr" rid="B151">2022</xref>)</td>
<td valign="top" align="left">ResNet-34</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">78.10</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">MLF (Feng L. et al., <xref ref-type="bibr" rid="B36">2022</xref>)</td>
<td valign="top" align="left">DS ResNet</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">70.36</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">LTMD (Wang S. et al., <xref ref-type="bibr" rid="B129">2022</xref>)</td>
<td valign="top" align="left">DenseNet</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">73.30</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">KLIF (Jiang and Zhang, <xref ref-type="bibr" rid="B55">2023</xref>)</td>
<td valign="top" align="left">CNN</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">70.90</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">MT-SNN (Wang X. et al., <xref ref-type="bibr" rid="B132">2023</xref>)</td>
<td valign="top" align="left">MT-VGG9</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">76.30</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">PSN (Fang et al., <xref ref-type="bibr" rid="B34">2023b</xref>)</td>
<td valign="top" align="left">VGG</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">85.90</td>
<td valign="top" align="left">Parallel</td>
</tr> <tr>
<td valign="top" align="left">PLIF (Fang et al., <xref ref-type="bibr" rid="B33">2021b</xref>)</td>
<td valign="top" align="left">PLIF-Net</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">97.57</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">MLF (Feng L. et al., <xref ref-type="bibr" rid="B36">2022</xref>)</td>
<td valign="top" align="left">DS ResNet</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">97.29</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left">KLIF (Jiang and Zhang, <xref ref-type="bibr" rid="B55">2023</xref>)</td>
<td valign="top" align="left">CNN</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">94.10</td>
<td valign="top" align="left">Time dependent</td>
</tr> <tr>
<td valign="top" align="left" rowspan="2">LSNN (Bellec et al., <xref ref-type="bibr" rid="B7">2018</xref>)</td>
<td valign="top" align="left" rowspan="2">LSTM</td>
<td valign="top" align="left">Sequential</td>
<td valign="top" align="center" rowspan="2">96.40</td>
<td valign="top" align="left" rowspan="2">Time dependent</td>
</tr>
<tr>
<td valign="top" align="left">MNIST</td>
</tr> <tr>
<td valign="top" align="left">ASN (Yin et al., <xref ref-type="bibr" rid="B156">2020</xref>)</td>
<td valign="top" align="left">RNN</td>
<td valign="top" align="left">PS-MNIST</td>
<td valign="top" align="center">97.90</td>
<td valign="top" align="left">Time dependent</td>
</tr>
<tr>
<td valign="top" align="left">SPSN (Yarga and Wood, <xref ref-type="bibr" rid="B153">2023</xref>)</td>
<td valign="top" align="left">MLP</td>
<td valign="top" align="left">SHD</td>
<td valign="top" align="center">86.89</td>
<td valign="top" align="left">Parallel</td>
</tr></tbody>
</table>
</table-wrap>
<sec>
<title>2.1 Spiking neurons with trainable parameters</title>
<p>Based on LIF, many improved spiking neuron models with trainable parameters have been proposed, which expand the representation space of neurons through parameter learning and improve the expression ability of SNNs. Fang et al. proposed Parametric LIF (PLIF) (Fang et al., <xref ref-type="bibr" rid="B33">2021b</xref>) by using trainable membrane time constant as follows:</p>
<disp-formula id="E4"><label>(4)</label><mml:math id="M4"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>H</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mi>V</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi><mml:mo>-</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>&#x0002B;</mml:mo><mml:mi>k</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>a</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>X</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>-</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>V</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi><mml:mo>-</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>-</mml:mo><mml:msub><mml:mrow><mml:mi>V</mml:mi></mml:mrow><mml:mrow><mml:mi>r</mml:mi><mml:mi>e</mml:mi><mml:mi>s</mml:mi><mml:mi>e</mml:mi><mml:mi>t</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>where <italic>k</italic>(<italic>a</italic>) in <xref ref-type="disp-formula" rid="E4">Equation (4)</xref> denotes a clamp function and <inline-formula><mml:math id="M5"><mml:mi>k</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>a</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mn>1</mml:mn><mml:mo>&#x0002B;</mml:mo><mml:mi>e</mml:mi><mml:mi>x</mml:mi><mml:mi>p</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mo>-</mml:mo><mml:mi>a</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow></mml:mfrac><mml:mo>&#x02208;</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mn>0</mml:mn><mml:mo>,</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:math></inline-formula>. The trainable membrane-related parameter of PLIF is biologically plausible, as neurons in the brain are heterogeneous. LTMD (Wang S. et al., <xref ref-type="bibr" rid="B129">2022</xref>) also leverages this biological plausibility but approaches it differently by employing learnable firing thresholds. An increase in the threshold of LTMD results in a reduction of output spikes, making an SNN less sensitive to its input and thus more robust. On the contrary, a decrease in the threshold leads to an increment of output spikes, making an SNN more sensitive to its input, which is particularly beneficial for processing transient small signals. Therefore, the learnable threshold <italic>V</italic><sub><italic>th</italic></sub> &#x0003D; tank(<italic>k</italic>), of which <italic>k</italic> is trainable, can lead to the optimal sensitivity of an SNN.</p>
<p>Diet-SNN (Rathi and Roy, <xref ref-type="bibr" rid="B107">2023</xref>) adopts an end-to-end gradient descent optimization algorithm to train the membrane-related parameters and firing thresholds of LIF neurons while optimizing the network weights. The trained neuron parameters selectively reduce the membrane potential, making spikes in the network sparser, thereby improving the computational efficiency of SNN. Spiking neurons with dynamic thresholds are adopted in LSNN (Bellec et al., <xref ref-type="bibr" rid="B7">2018</xref>). After firing a spike each time, the firing threshold of a neuron will increase by a fixed amount, and then it will decay exponentially according to the time constant. Adaptive spiking neuron (ASN) (Yin et al., <xref ref-type="bibr" rid="B156">2020</xref>) was proposed for sequence and streaming media tasks. In ASN, the time constant of membrane potential is trainable. In addition, similar to LSNN, the firing threshold will increase after each spike of the neuron, thus improving sparsity and efficiency.</p>
<p>In KLIF (Jiang and Zhang, <xref ref-type="bibr" rid="B55">2023</xref>), a trainable scaling factor <italic>k</italic> and a nonlinear ReLU activation function are inserted between charging and firing. The dynamics of KLIF can be described by <xref ref-type="disp-formula" rid="E1">Equations (1)</xref>, (<xref ref-type="disp-formula" rid="E5">5</xref>&#x02013;<xref ref-type="disp-formula" rid="E7">7</xref>).</p>
<disp-formula id="E5"><label>(5)</label><mml:math id="M6"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>F</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mtext>ReLU</mml:mtext><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>k</mml:mi><mml:mi>H</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E6"><label>(6)</label><mml:math id="M7"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>S</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mo>&#x00398;</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>F</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>-</mml:mo><mml:msub><mml:mrow><mml:mi>V</mml:mi></mml:mrow><mml:mrow><mml:mi>t</mml:mi><mml:mi>h</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E7"><label>(7)</label><mml:math id="M8"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>V</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mi>F</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mn>1</mml:mn><mml:mo>-</mml:mo><mml:mi>S</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>&#x0002B;</mml:mo><mml:msub><mml:mrow><mml:mi>V</mml:mi></mml:mrow><mml:mrow><mml:mi>r</mml:mi><mml:mi>e</mml:mi><mml:mi>s</mml:mi><mml:mi>e</mml:mi><mml:mi>t</mml:mi></mml:mrow></mml:msub><mml:mi>S</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>.</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>Compared with LIF, KLIF can automatically adjust the membrane potential and the gradient of backpropagation within the neuron. GLIF (Yao et al., <xref ref-type="bibr" rid="B151">2022</xref>) introduces a gating unit that fuses multiple biometric features, with the ratio of these features adjusted by a trainable gating factor. Moreover, inspired by various spiking patterns of brain neurons, LIFB (Shen et al., <xref ref-type="bibr" rid="B116">2023</xref>) has three modes: resting, regular spiking, and burst spiking. The density of the burst spiking can be learned automatically, which greatly enriches the representation capability of neurons.</p>
<p>In addition, there are other studies trying to improve performance by multi-level firing thresholds instead of trainable parameters. To reduce the performance loss caused by the transmission of binarized spikes in the network, MT-SNN (Wang X. et al., <xref ref-type="bibr" rid="B132">2023</xref>) introduces multi-level firing thresholds. MT-SNN performs convolution operations on the binarized spikes generated by different firing thresholds and then sums them up. Similarly, MLF (Feng L. et al., <xref ref-type="bibr" rid="B36">2022</xref>) can also fire spikes under different firing thresholds, thus improving the performance of SNNs.</p>
</sec>
<sec>
<title>2.2 Parallel spiking neurons</title>
<p>A typical neuron model like LIF is time-dependent, that is, its state at time <italic>t</italic> relies on its state at time <italic>t</italic> &#x02212; 1, resulting in a high computation load. Fang et al. (<xref ref-type="bibr" rid="B34">2023b</xref>) proposed a parallel spiking neuron (PSN) to accelerate the computation by parallel computing. By eliminating the resetting process, they represent the charging process of PSN by a non-iterative equation as <xref ref-type="disp-formula" rid="E8">Equation (8)</xref>:</p>
<disp-formula id="E8"><label>(8)</label><mml:math id="M9"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>H</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:munderover accentunder="false" accent="false"><mml:mrow><mml:mo>&#x02211;</mml:mo></mml:mrow><mml:mrow><mml:mi>i</mml:mi><mml:mo>=</mml:mo><mml:mn>0</mml:mn></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mo>-</mml:mo><mml:mn>1</mml:mn></mml:mrow></mml:munderover></mml:mstyle><mml:msub><mml:mrow><mml:mi>W</mml:mi></mml:mrow><mml:mrow><mml:mi>t</mml:mi><mml:mo>,</mml:mo><mml:mi>i</mml:mi></mml:mrow></mml:msub><mml:mo>&#x000B7;</mml:mo><mml:mi>X</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>i</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>where <italic>W</italic><sub><italic>t,i</italic></sub> is the weight between input <italic>X</italic>[<italic>i</italic>] and membrane potential <italic>H</italic>[<italic>t</italic>]. For LIF neuron, <inline-formula><mml:math id="M10"><mml:msub><mml:mrow><mml:mi>W</mml:mi></mml:mrow><mml:mrow><mml:mi>t</mml:mi><mml:mo>,</mml:mo><mml:mi>i</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:msub><mml:mrow><mml:mi>&#x003C4;</mml:mi></mml:mrow><mml:mrow><mml:mi>m</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:mfrac><mml:msup><mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mn>1</mml:mn><mml:mo>-</mml:mo><mml:mfrac><mml:mrow><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:msub><mml:mrow><mml:mi>&#x003C4;</mml:mi></mml:mrow><mml:mrow><mml:mi>m</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:mfrac></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mi>t</mml:mi><mml:mo>-</mml:mo><mml:mi>i</mml:mi></mml:mrow></mml:msup></mml:math></inline-formula>. The dynamics of PSN are as <xref ref-type="disp-formula" rid="E9">Equations (9</xref>, <xref ref-type="disp-formula" rid="E10">10)</xref>:</p>
<disp-formula id="E9"><label>(9)</label><mml:math id="M11"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mstyle mathvariant="bold-italic"><mml:mi>H</mml:mi></mml:mstyle><mml:mo>=</mml:mo><mml:mstyle mathvariant="bold-italic"><mml:mi>W</mml:mi></mml:mstyle><mml:mstyle mathvariant="bold-italic"><mml:mi>X</mml:mi></mml:mstyle><mml:mo>,</mml:mo><mml:mtext>&#x000A0;&#x000A0;</mml:mtext><mml:mstyle mathvariant="bold-italic"><mml:mi>W</mml:mi></mml:mstyle><mml:mo>&#x02208;</mml:mo><mml:msup><mml:mrow><mml:mi>&#x0211D;</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mo>&#x000D7;</mml:mo><mml:mi>T</mml:mi></mml:mrow></mml:msup><mml:mo>,</mml:mo><mml:mstyle mathvariant="bold-italic"><mml:mi>X</mml:mi></mml:mstyle><mml:mo>&#x02208;</mml:mo><mml:msup><mml:mrow><mml:mi>&#x0211D;</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mo>&#x000D7;</mml:mo><mml:mi>N</mml:mi></mml:mrow></mml:msup></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E10"><label>(10)</label><mml:math id="M12"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mstyle mathvariant="bold-italic"><mml:mi>S</mml:mi></mml:mstyle><mml:mo>=</mml:mo><mml:mo>&#x00398;</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mstyle mathvariant="bold-italic"><mml:mi>H</mml:mi></mml:mstyle><mml:mo>-</mml:mo><mml:mstyle mathvariant="bold-italic"><mml:mi>B</mml:mi></mml:mstyle></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo><mml:mtext>&#x000A0;&#x000A0;</mml:mtext><mml:mstyle mathvariant="bold-italic"><mml:mi>B</mml:mi></mml:mstyle><mml:mo>&#x02208;</mml:mo><mml:msup><mml:mrow><mml:mi>&#x0211D;</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi></mml:mrow></mml:msup><mml:mo>,</mml:mo><mml:mstyle mathvariant="bold-italic"><mml:mi>S</mml:mi></mml:mstyle><mml:mo>&#x02208;</mml:mo><mml:msup><mml:mrow><mml:mrow><mml:mo>{</mml:mo><mml:mrow><mml:mn>0</mml:mn><mml:mo>,</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mo>}</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mo>&#x000D7;</mml:mo><mml:mi>N</mml:mi></mml:mrow></mml:msup></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>where <italic><bold>X</bold></italic> is the input, <italic><bold>W</bold></italic> and <italic><bold>B</bold></italic> are trainable weights and trainable firing thresholds, respectively. <italic><bold>H</bold></italic> is the membrane potential after charging, and <italic><bold>S</bold></italic> denotes whether a neuron spikes. <italic>N</italic> and <italic>T</italic> are the batch size and the number of time steps, respectively. For step-by-step serial forward computation and variable-length sequence processing, the masked PSN and the sliding PSN are also derived.</p>
<p>The stochastic parallel spiking neuron (SPSN) (Yarga and Wood, <xref ref-type="bibr" rid="B153">2023</xref>) adopts an idea similar to PSN, by removing the resetting mechanism. The neuronal dynamics of SPSN contains two parts, namely parallel leaky integrator and stochastic firing. The leaky integrator is a linear time-invariant system, which can be transformed into the Fourier domain to realize parallel computation. Stochastic firing adaptively adjusts the firing probability through trainable parameters, enhancing the network&#x00027;s capability to process information in a dynamic and efficient manner.</p>
</sec>
</sec>
<sec id="s3">
<title>3 Fundamentals of spiking neural networks</title>
<sec>
<title>3.1 Information coding</title>
<p>To process image data through SNNs, it is essential to first encode the data into spike trains. Rate coding (Adrian and Zotterman, <xref ref-type="bibr" rid="B2">1926</xref>) is the most commonly used information coding method in SNNs, in which the firing rate is proportional to the intensity of the input signal and spikes are typically generated by a Poisson process (Wiener and Richmond, <xref ref-type="bibr" rid="B137">2003</xref>). To encode information more accurately, rate coding requires a longer time window, which leads to a slower information transmission rate. In contrast, utilizing a shorter time window may result in loss of information during encoding, presenting a trade-off between speed and accuracy in information transmission.</p>
<p>Different from rate coding, temporal coding represents information through the timing of spikes. Time-to-first-spike (TTFS) (Park et al., <xref ref-type="bibr" rid="B98">2020</xref>; Guo W. et al., <xref ref-type="bibr" rid="B42">2021</xref>) stands out for its simplicity and efficiency in temporal coding, which uses the time of the first spike fired by the neuron to represent the input signal. TTFS effectively reduces the total number of spikes, thereby accelerating the computation of SNNs. TTFS algorithm can be described as <xref ref-type="disp-formula" rid="E11">Equation (11)</xref>:</p>
<disp-formula id="E11"><label>(11)</label><mml:math id="M13"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>S</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mo>]</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mrow><mml:mo>{</mml:mo><mml:mrow><mml:mtable style="text-align:axis;" equalrows="false" columnlines="none" equalcolumns="false" class="array"><mml:mtr><mml:mtd><mml:mn>1</mml:mn><mml:mo>,</mml:mo></mml:mtd><mml:mtd><mml:mtext class="textrm" mathvariant="normal">if</mml:mtext><mml:mi>t</mml:mi><mml:mo>=</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mfrac><mml:mrow><mml:msub><mml:mrow><mml:mi>X</mml:mi></mml:mrow><mml:mrow><mml:mi>m</mml:mi><mml:mi>a</mml:mi><mml:mi>x</mml:mi></mml:mrow></mml:msub><mml:mo>-</mml:mo><mml:mi>X</mml:mi></mml:mrow><mml:mrow><mml:msub><mml:mrow><mml:mi>X</mml:mi></mml:mrow><mml:mrow><mml:mi>m</mml:mi><mml:mi>a</mml:mi><mml:mi>x</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:mfrac></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:msub><mml:mrow><mml:mi>t</mml:mi></mml:mrow><mml:mrow><mml:mi>m</mml:mi><mml:mi>a</mml:mi><mml:mi>x</mml:mi></mml:mrow></mml:msub></mml:mtd></mml:mtr><mml:mtr><mml:mtd><mml:mn>0</mml:mn><mml:mo>,</mml:mo></mml:mtd><mml:mtd><mml:mtext class="textrm" mathvariant="normal">otherwise</mml:mtext></mml:mtd></mml:mtr></mml:mtable></mml:mrow></mml:mrow><mml:mtext>&#x000A0;</mml:mtext><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>where <italic>S</italic>[<italic>t</italic>] represents whether a spike is fired at time <italic>t</italic> after encoding, <italic>t</italic><sub><italic>max</italic></sub> denotes the maximum time allowed during encoding, <italic>X</italic> and <italic>X</italic><sub><italic>max</italic></sub> represent the input signal and its maximum value, respectively. In the TTFS encoding method, larger values of the input signal lead to earlier firing of spikes.</p>
</sec>
<sec>
<title>3.2 Network training</title>
<sec>
<title>3.2.1 Surrogate gradient</title>
<p>As the core components of SNNs, neurons are essential for information processing and transmission, since spikes are fired by neurons. However, the firing of spikes involves the non-differentiable Heaviside step function, which presents a significant challenge in the direct training of SNNs. To address the non-differentiability of the Heaviside step function, Neftci et al. (<xref ref-type="bibr" rid="B94">2019</xref>) proposed the Surrogate Gradient (SG) algorithm. In SG, the Heaviside step function is adopted to generate spikes during forward propagation, and differentiable functions are adopted for gradient calculation during backpropagation. Notably, SG functions could vary according to the networks. For instance, the SG function used in SEW ResNet (Fang et al., <xref ref-type="bibr" rid="B35">2021a</xref>) is the derivative of the arctan function as follows:</p>
<disp-formula id="E12"><label>(12)</label><mml:math id="M14"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>&#x003C3;</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>x</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>&#x003C0;</mml:mi></mml:mrow></mml:mfrac><mml:mo class="qopname">arctan</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mfrac><mml:mrow><mml:mi>&#x003C0;</mml:mi></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:mfrac><mml:mi>&#x003B1;</mml:mi><mml:mi>x</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>&#x0002B;</mml:mo><mml:mfrac><mml:mrow><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:mfrac><mml:mtext>&#x000A0;</mml:mtext><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E13"><label>(13)</label><mml:math id="M15"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:msup><mml:mrow><mml:mi>&#x003C3;</mml:mi></mml:mrow><mml:mrow><mml:mi>&#x02032;</mml:mi></mml:mrow></mml:msup><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>x</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>&#x003B1;</mml:mi></mml:mrow><mml:mrow><mml:mn>2</mml:mn><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mn>1</mml:mn><mml:mo>&#x0002B;</mml:mo><mml:msup><mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mfrac><mml:mrow><mml:mi>&#x003C0;</mml:mi></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:mfrac><mml:mi>&#x003B1;</mml:mi><mml:mi>x</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:msup></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow></mml:mfrac><mml:mtext>&#x000A0;</mml:mtext><mml:mo>.</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p><xref ref-type="disp-formula" rid="E13">Equation (13)</xref> is the derivative of <xref ref-type="disp-formula" rid="E12">Equation (12)</xref>. In addition, SG could be the derivative of Sigmoid (<xref ref-type="fig" rid="F1">Figure 1B</xref>) (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>,<xref ref-type="bibr" rid="B171">b</xref>; Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>), tanh (Guo et al., <xref ref-type="bibr" rid="B43">2022a</xref>), or rectangular (Wu et al., <xref ref-type="bibr" rid="B140">2018</xref>, <xref ref-type="bibr" rid="B141">2019</xref>) functions, etc. To address the problem of gradient vanishing caused by a surrogate gradient function with fixed parameters, Lian et al. (<xref ref-type="bibr" rid="B75">2023</xref>) proposed the Learnable Surrogate Gradient (LSG), in which a learnable parameter is used to adjust the gradient-available interval.</p>
<p>Li et al. (<xref ref-type="bibr" rid="B73">2021</xref>) proposed Differentiable Spike (Dspike) as another approach to overcome the non-differentiable problem of the Heaviside function. Based on the hyperbolic tangent function, Dspike can be described as <xref ref-type="disp-formula" rid="E14">Equation (14)</xref>:</p>
<disp-formula id="E14"><label>(14)</label><mml:math id="M16"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>D</mml:mi><mml:mi>s</mml:mi><mml:mi>p</mml:mi><mml:mi>i</mml:mi><mml:mi>k</mml:mi><mml:mi>e</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>x</mml:mi><mml:mo>,</mml:mo><mml:mi>b</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mo class="qopname">tanh</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>b</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>x</mml:mi><mml:mo>-</mml:mo><mml:mn>0</mml:mn><mml:mo>.</mml:mo><mml:mn>5</mml:mn></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>&#x0002B;</mml:mo><mml:mo class="qopname">tanh</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>b</mml:mi><mml:mo>/</mml:mo><mml:mn>2</mml:mn></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mn>2</mml:mn><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mo class="qopname">tanh</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>b</mml:mi><mml:mo>/</mml:mo><mml:mn>2</mml:mn></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow></mml:mfrac><mml:mo>,</mml:mo><mml:mtext>if&#x000A0;</mml:mtext><mml:mn>0</mml:mn><mml:mo>&#x02264;</mml:mo><mml:mi>x</mml:mi><mml:mo>&#x02264;</mml:mo><mml:mn>1</mml:mn></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>By adjusting the parameter <italic>b</italic>, different backpropagation gradients can be obtained. Differentiation on Spike Representation (DSR) proposed by Meng et al. (<xref ref-type="bibr" rid="B89">2022</xref>) encodes spike trains and represents them as sub-differentiable mapping, which also avoids the non-differentiable problem during backpropagation.</p>
</sec>
<sec>
<title>3.2.2 Loss function and backpropagation</title>
<p>Loss function is the key to neural network training, and different loss functions have been proposed to enhance the performance of SNNs. IM-Loss (Guo et al., <xref ref-type="bibr" rid="B43">2022a</xref>), for example, aims to maximize the information flow in the network. The total loss function consists of two parts, cross-entropy loss, and IM-Loss, as <xref ref-type="disp-formula" rid="E15">Equations (15</xref>, <xref ref-type="disp-formula" rid="E16">16)</xref>:</p>
<disp-formula id="E15"><label>(15)</label><mml:math id="M17"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:msub><mml:mrow><mml:mrow><mml:mi mathvariant="script">L</mml:mi></mml:mrow></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>o</mml:mi><mml:mi>t</mml:mi><mml:mi>a</mml:mi><mml:mi>l</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:msub><mml:mrow><mml:mrow><mml:mi mathvariant="script">L</mml:mi></mml:mrow></mml:mrow><mml:mrow><mml:mi>C</mml:mi><mml:mi>E</mml:mi></mml:mrow></mml:msub><mml:mo>&#x0002B;</mml:mo><mml:mi>&#x003BB;</mml:mi><mml:msub><mml:mrow><mml:mrow><mml:mi mathvariant="script">L</mml:mi></mml:mrow></mml:mrow><mml:mrow><mml:mi>I</mml:mi><mml:mi>M</mml:mi></mml:mrow></mml:msub><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E16"><label>(16)</label><mml:math id="M18"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:msub><mml:mrow><mml:mrow><mml:mi mathvariant="script">L</mml:mi></mml:mrow></mml:mrow><mml:mrow><mml:mi>I</mml:mi><mml:mi>M</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:munderover accentunder="false" accent="false"><mml:mrow><mml:mo>&#x02211;</mml:mo></mml:mrow><mml:mrow><mml:mi>l</mml:mi><mml:mo>=</mml:mo><mml:mn>0</mml:mn></mml:mrow><mml:mrow><mml:mi>L</mml:mi></mml:mrow></mml:munderover></mml:mstyle><mml:msup><mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:msub><mml:mrow><mml:mi>&#x0016A;</mml:mi></mml:mrow><mml:mrow><mml:mi>l</mml:mi></mml:mrow></mml:msub><mml:mo>-</mml:mo><mml:msub><mml:mrow><mml:mi>V</mml:mi></mml:mrow><mml:mrow><mml:mi>t</mml:mi><mml:mi>h</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:msup><mml:mo>/</mml:mo><mml:mi>L</mml:mi><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>where &#x0016A;<sub><italic>l</italic></sub> is the averaged membrane potential at all time steps of the <italic>l</italic>-th layer, and <italic>L</italic> is the total number of layers. To alleviate the information loss in SNNs and reduce the quantization error, RMP-Loss (Guo et al., <xref ref-type="bibr" rid="B46">2023a</xref>) is proposed to adjust the distribution of membrane potential. RecDis-SNN (Guo et al., <xref ref-type="bibr" rid="B44">2022b</xref>) adopts MDP-Loss that also adjusts the membrane potential distribution to overcome the distribution shift during network training. In addition, to improve the generalization ability of SNNs, Deng et al. proposed temporal efficient training (TET) (Deng et al., <xref ref-type="bibr" rid="B27">2022</xref>) loss function to make the network output closer to the target distribution.</p>
<p>Distinct from ANNs, there&#x00027;s an additional dimension in SNNs, the temporal domain. For spiking neurons, the membrane potential in the current step depends on the membrane potential in the previous time step, that is, there is a time dependence. Thus, backpropagation in ANNs does not apply to SNNs. Backpropagation Through Time (BPTT) (Werbos, <xref ref-type="bibr" rid="B136">1990</xref>; Bird and Polivoda, <xref ref-type="bibr" rid="B11">2021</xref>), originally developed for recurrent neural networks (RNNs), is applied to SNNs due to their similar characteristics to those of RNNs. The combination of BPTT and surrogate gradient is the basic approach in SNNs. Spatio-temporal backpropagation (STBP) (Wu et al., <xref ref-type="bibr" rid="B140">2018</xref>), proposed by Wu et al., takes the gradient update in both the spatial domain and temporal domain into account to train SNNs. However, the additional time dimension exposes BPTT and STBP to the problem of requiring a large amount of training memory and training time. Therefore, Xiao et al. (<xref ref-type="bibr" rid="B143">2022</xref>) proposed an online training through time (OTTT) algorithm derived from BPTT, which only requires constant training memory consumption agnostic to time steps, and reduces the significant memory costs compared to BPTT. The backward of BPTT and OTTT are shown in <xref ref-type="fig" rid="F2">Figure 2</xref>. Another efficient backpropagation method, Spatial Learning Through Time (SLTT) (Meng et al., <xref ref-type="bibr" rid="B90">2023</xref>), ignores the unimportant routes in the computational graph during backpropagation, to reduce training memory consumption and training time. However, although OTTT and SLTT show better training memory consumption than BPTT, direct training high-performance SNNs are still dominated by the combination of BPTT and surrogate gradient, such as SGLFormer (Zhang et al., <xref ref-type="bibr" rid="B161">2024</xref>), Spikformer (Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>), etc. Thus, it&#x00027;s essential to investigate direct training methods offering both high effectiveness and efficiency.</p>
<fig id="F2" position="float">
<label>Figure 2</label>
<caption><p>The backward of <bold>(A)</bold> BPTT and <bold>(B)</bold> OTTT.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnins-18-1383844-g0002.tif"/>
</fig>
</sec>
<sec>
<title>3.2.3 Batch normalization</title>
<p>In SNNs, batch normalization is an indispensable component, especially in the context that deep SNNs are difficult to train and converge, compared to ANNs. To mitigate the degradation problems of SNNs, Zheng et al. (<xref ref-type="bibr" rid="B168">2021</xref>) proposed threshold-dependent batch normalization (tdBN), which is described as <xref ref-type="disp-formula" rid="E17">Equation (17)</xref>:</p>
<disp-formula id="E17"><label>(17)</label><mml:math id="M19"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:msub><mml:mrow><mml:mover accent="true"><mml:mrow><mml:mi>X</mml:mi></mml:mrow><mml:mo>^</mml:mo></mml:mover></mml:mrow><mml:mrow><mml:mi>k</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:msub><mml:mrow><mml:mi>&#x003B3;</mml:mi></mml:mrow><mml:mrow><mml:mi>k</mml:mi></mml:mrow></mml:msub><mml:mfrac><mml:mrow><mml:mi>&#x003B1;</mml:mi><mml:msub><mml:mrow><mml:mi>V</mml:mi></mml:mrow><mml:mrow><mml:mi>t</mml:mi><mml:mi>h</mml:mi></mml:mrow></mml:msub><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:msub><mml:mrow><mml:mi>X</mml:mi></mml:mrow><mml:mrow><mml:mi>k</mml:mi></mml:mrow></mml:msub><mml:mo>-</mml:mo><mml:mi>&#x003BC;</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:msqrt><mml:mrow><mml:msup><mml:mrow><mml:mi>&#x003C3;</mml:mi></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:msup><mml:mo>&#x0002B;</mml:mo><mml:mi>&#x003F5;</mml:mi></mml:mrow></mml:msqrt></mml:mrow></mml:mfrac><mml:mo>&#x0002B;</mml:mo><mml:msub><mml:mrow><mml:mi>&#x003B2;</mml:mi></mml:mrow><mml:mrow><mml:mi>k</mml:mi></mml:mrow></mml:msub><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>where &#x003B1; is a hyperparameter, <italic>V</italic><sub><italic>th</italic></sub> is the firing threshold of the neuron, <italic>X</italic><sub><italic>k</italic></sub> is the feature of the <italic>k</italic>-th channel, &#x003B3;<sub><italic>k</italic></sub> and &#x003B2;<sub><italic>k</italic></sub> are trainable parameters, &#x003BC; and &#x003C3;<sup>2</sup> are mean and variance, respectively, &#x003F5; is a tiny constant. Temporal effective batch normalization (TEBN) (Duan et al., <xref ref-type="bibr" rid="B29">2022</xref>) regularizes the temporal distribution, by adopting batch normalization with different parameters at different time steps. Batch normalization through time (BNTT) proposed by Kim and Panda (<xref ref-type="bibr" rid="B59">2021</xref>) is similar to TEBN, which also adopts different batch normalization parameters for feature maps at different time steps. Moreover, Guo et al. (<xref ref-type="bibr" rid="B45">2023b</xref>) applied batch normalization inside the LIF neuron to normalize the distribution of membrane potentials before firing spikes.</p>
</sec>
</sec>
</sec>
<sec id="s4">
<title>4 SNN architecture developments</title>
<p>This review focuses on the most recent SNN models. Recently, the evolution of residual blocks enhances both the size and performance of deep SNNs significantly. In addition, combining SNNs with transformer architecture has broken the bottleneck of SNNs&#x00027; performance. Therefore, this review focuses on the application of two kinds of architectures in direct training deep SNNs: transformer structures (Section 4.1) and the residual connections (Section 4.2). <xref ref-type="table" rid="T2">Table 2</xref> summarizes their performance on mainstream datasets (ImageNet-1K, CIFAR10, CIFAR100, DVS128 Gesture, CIFAR10-DVS).</p>
<table-wrap position="float" id="T2">
<label>Table 2</label>
<caption><p>Overview of direct training deep SNNs and their performance on ImageNet, CIFAR10, CIFAR100, DVS128-Gesture, CIFAR10-DVS.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>Method</bold></th>
<th valign="top" align="left"><bold>Architecture</bold></th>
<th valign="top" align="center"><bold>Param (M)</bold></th>
<th valign="top" align="center"><bold>Time steps</bold></th>
<th valign="top" align="left"><bold>Dataset</bold></th>
<th valign="top" align="center"><bold>Top-1 Acc (<italic>%</italic>)</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Spiking ResNet (Hu et al., <xref ref-type="bibr" rid="B50">2021a</xref>)</td>
<td valign="top" align="left">ResNet-50</td>
<td valign="top" align="center">25.56</td>
<td valign="top" align="center">350</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">72.75</td>
</tr> <tr>
<td valign="top" align="left">SEW ResNet (Fang et al., <xref ref-type="bibr" rid="B35">2021a</xref>)</td>
<td valign="top" align="left">SEW-ResNet-152</td>
<td valign="top" align="center">60.19</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">69.26</td>
</tr> <tr>
<td valign="top" align="left">MS-ResNet (Hu et al., <xref ref-type="bibr" rid="B51">2021b</xref>)</td>
<td valign="top" align="left">MS-ResNet-104</td>
<td valign="top" align="center">77.28</td>
<td valign="top" align="center">5</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">76.02</td>
</tr> <tr>
<td valign="top" align="left">Att MS-ResNet (Yao et al., <xref ref-type="bibr" rid="B150">2023b</xref>)</td>
<td valign="top" align="left">Att-MS-ResNet-104</td>
<td valign="top" align="center">78.37</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">77.08</td>
</tr> <tr>
<td valign="top" align="left">Spikformer (Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>)</td>
<td valign="top" align="left">Spikformer-8-768</td>
<td valign="top" align="center">66.34</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">74.81</td>
</tr> <tr>
<td valign="top" align="left">Spikingformer (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>)</td>
<td valign="top" align="left">Spikingformer-8-768</td>
<td valign="top" align="center">66.34</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">75.85</td>
</tr> <tr>
<td valign="top" align="left">CML (Zhou et al., <xref ref-type="bibr" rid="B171">2023b</xref>)</td>
<td valign="top" align="left">Spikformer-8-768</td>
<td valign="top" align="center">66.34</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">77.34</td>
</tr> <tr>
<td valign="top" align="left">Spike-driven Transformer (Yao et al., <xref ref-type="bibr" rid="B152">2023a</xref>)</td>
<td valign="top" align="left">S-Transformer-8-768</td>
<td valign="top" align="center">66.34</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">77.07</td>
</tr> <tr>
<td valign="top" align="left">SpikingResformer (Shi et al., <xref ref-type="bibr" rid="B118">2024</xref>)</td>
<td valign="top" align="left">SpikingResformer-L</td>
<td valign="top" align="center">60.38</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">79.40</td>
</tr> <tr>
<td valign="top" align="left">Spike-driven Transformer V2 (Yao et al., <xref ref-type="bibr" rid="B149">2024</xref>)</td>
<td valign="top" align="left">Meta-SpikeFormer</td>
<td valign="top" align="center">55.40</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">80.00</td>
</tr> <tr>
<td valign="top" align="left">Spikformer V2 (Zhou Z. et al., <xref ref-type="bibr" rid="B173">2024</xref>)</td>
<td valign="top" align="left">Spikformer V2-8-512</td>
<td valign="top" align="center">51.55</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">80.38</td>
</tr> <tr>
<td valign="top" align="left">SGLFormer (Zhang et al., <xref ref-type="bibr" rid="B161">2024</xref>)</td>
<td valign="top" align="left">SGLFormer-8-768</td>
<td valign="top" align="center">64.02</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">83.73</td>
</tr> <tr>
<td valign="top" align="left">QKFormer (Zhou C. et al., <xref ref-type="bibr" rid="B170">2024</xref>)</td>
<td valign="top" align="left">HST-10-768</td>
<td valign="top" align="center">64.96</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">ImageNet</td>
<td valign="top" align="center">85.65</td>
</tr> <tr>
<td valign="top" align="left">Hybrid training (Rathi et al., <xref ref-type="bibr" rid="B108">2020</xref>)</td>
<td valign="top" align="left">VGG-11</td>
<td valign="top" align="center">9.27</td>
<td valign="top" align="center">125</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">92.22</td>
</tr> <tr>
<td valign="top" align="left">STBP-tdBN (Zheng et al., <xref ref-type="bibr" rid="B168">2021</xref>)</td>
<td valign="top" align="left">ResNet-19</td>
<td valign="top" align="center">12.63</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">92.92</td>
</tr> <tr>
<td valign="top" align="left">TET (Deng et al., <xref ref-type="bibr" rid="B27">2022</xref>)</td>
<td valign="top" align="left">ResNet-19</td>
<td valign="top" align="center">12.63</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">94.44</td>
</tr> <tr>
<td valign="top" align="left">MS-ResNet (Hu et al., <xref ref-type="bibr" rid="B51">2021b</xref>)</td>
<td valign="top" align="left">MS-ResNet-110</td>
<td valign="top" align="center">&#x02013;</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">92.12</td>
</tr> <tr>
<td valign="top" align="left">Spikformer (Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>)</td>
<td valign="top" align="left">Spikformer-4-384</td>
<td valign="top" align="center">9.32</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">95.51</td>
</tr> <tr>
<td valign="top" align="left">Spikingformer (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>)</td>
<td valign="top" align="left">Spikingformer-4-384</td>
<td valign="top" align="center">9.32</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">95.81</td>
</tr> <tr>
<td valign="top" align="left">CML (Zhou et al., <xref ref-type="bibr" rid="B171">2023b</xref>)</td>
<td valign="top" align="left">Spikformer-4-384</td>
<td valign="top" align="center">9.32</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">96.04</td>
</tr> <tr>
<td valign="top" align="left">Spike-driven Transformer (Yao et al., <xref ref-type="bibr" rid="B152">2023a</xref>)</td>
<td valign="top" align="left">S-Transformer-2-512</td>
<td valign="top" align="center">10.23</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">95.60</td>
</tr> <tr>
<td valign="top" align="left">SGLFormer (Zhang et al., <xref ref-type="bibr" rid="B161">2024</xref>)</td>
<td valign="top" align="left">SGLFormer-4-384</td>
<td valign="top" align="center">8.85</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR10</td>
<td valign="top" align="center">96.76</td>
</tr> <tr>
<td valign="top" align="left">Hybrid training (Rathi et al., <xref ref-type="bibr" rid="B108">2020</xref>)</td>
<td valign="top" align="left">VGG-11</td>
<td valign="top" align="center">9.27</td>
<td valign="top" align="center">125</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">67.87</td>
</tr> <tr>
<td valign="top" align="left">STBP-tdBN (Zheng et al., <xref ref-type="bibr" rid="B168">2021</xref>)</td>
<td valign="top" align="left">ResNet-19</td>
<td valign="top" align="center">12.63</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">70.86</td>
</tr> <tr>
<td valign="top" align="left">TET (Deng et al., <xref ref-type="bibr" rid="B27">2022</xref>)</td>
<td valign="top" align="left">ResNet-19</td>
<td valign="top" align="center">12.63</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">74.47</td>
</tr> <tr>
<td valign="top" align="left">Spikformer (Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>)</td>
<td valign="top" align="left">Spikformer-4-384</td>
<td valign="top" align="center">9.32</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">78.21</td>
</tr> <tr>
<td valign="top" align="left">Spikingformer (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>)</td>
<td valign="top" align="left">Spikingformer-4-384</td>
<td valign="top" align="center">9.32</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">79.21</td>
</tr> <tr>
<td valign="top" align="left">CML (Zhou et al., <xref ref-type="bibr" rid="B171">2023b</xref>)</td>
<td valign="top" align="left">Spikformer-4-384</td>
<td valign="top" align="center">9.32</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">80.02</td>
</tr> <tr>
<td valign="top" align="left">Spike-driven Transformer (Yao et al., <xref ref-type="bibr" rid="B152">2023a</xref>)</td>
<td valign="top" align="left">S-Transformer-2-512</td>
<td valign="top" align="center">10.28</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">78.4</td>
</tr> <tr>
<td valign="top" align="left">SGLFormer (Zhang et al., <xref ref-type="bibr" rid="B161">2024</xref>)</td>
<td valign="top" align="left">SGLFormer-4-384</td>
<td valign="top" align="center">8.88</td>
<td valign="top" align="center">4</td>
<td valign="top" align="left">CIFAR100</td>
<td valign="top" align="center">82.26</td>
</tr>
<tr>
<td valign="top" align="left">SEW-ResNet (Hu et al., <xref ref-type="bibr" rid="B51">2021b</xref>)</td>
<td valign="top" align="left">SEW-ResNet</td>
<td valign="top" align="center">&#x02013;</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">97.9</td>
</tr> <tr>
<td valign="top" align="left">tdBN (Zheng et al., <xref ref-type="bibr" rid="B168">2021</xref>)</td>
<td valign="top" align="left">ResNet</td>
<td valign="top" align="center">&#x02013;</td>
<td valign="top" align="center">40</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">96.9</td>
</tr> <tr>
<td valign="top" align="left">Spikformer (Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>)</td>
<td valign="top" align="left">Spikformer-2-256</td>
<td valign="top" align="center">2.57</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">98.3</td>
</tr> <tr>
<td valign="top" align="left">Spikingformer (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>)</td>
<td valign="top" align="left">Spikingformer-2-256</td>
<td valign="top" align="center">2.57</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">98.3</td>
</tr> <tr>
<td valign="top" align="left">CML (Zhou et al., <xref ref-type="bibr" rid="B171">2023b</xref>)</td>
<td valign="top" align="left">Spikformer-2-256</td>
<td valign="top" align="center">2.57</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">98.6</td>
</tr> <tr>
<td valign="top" align="left">Spike-driven Transformer (Yao et al., <xref ref-type="bibr" rid="B152">2023a</xref>)</td>
<td valign="top" align="left">S-Transformer-2-256</td>
<td valign="top" align="center">2.57</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">99.3</td>
</tr> <tr>
<td valign="top" align="left">STSA (Wang Y. et al., <xref ref-type="bibr" rid="B134">2023</xref>)</td>
<td valign="top" align="left">STSFormer-2-256</td>
<td valign="top" align="center">1.99</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">98.72</td>
</tr>
<tr>
<td valign="top" align="left">SGLFormer (Zhang et al., <xref ref-type="bibr" rid="B161">2024</xref>)</td>
<td valign="top" align="left">SGLFormer-3-256</td>
<td valign="top" align="center">2.17</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">DVS128-Gesture</td>
<td valign="top" align="center">98.6</td>
</tr> <tr>
<td valign="top" align="left">SEW-ResNet (Hu et al., <xref ref-type="bibr" rid="B51">2021b</xref>)</td>
<td valign="top" align="left">SEW-ResNet</td>
<td valign="top" align="center">&#x02013;</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">74.4</td>
</tr> <tr>
<td valign="top" align="left">Spikformer (Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>)</td>
<td valign="top" align="left">Spikformer-2-256</td>
<td valign="top" align="center">2.57</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">80.9</td>
</tr> <tr>
<td valign="top" align="left">Spikingformer (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>)</td>
<td valign="top" align="left">Spikingformer-2-256</td>
<td valign="top" align="center">2.57</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">81.3</td>
</tr> <tr>
<td valign="top" align="left">CML (Zhou et al., <xref ref-type="bibr" rid="B171">2023b</xref>)</td>
<td valign="top" align="left">Spikformer-2-256</td>
<td valign="top" align="center">2.57</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">80.9</td>
</tr> <tr>
<td valign="top" align="left">Spike-driven Transformer (Yao et al., <xref ref-type="bibr" rid="B152">2023a</xref>)</td>
<td valign="top" align="left">S-Transformer-2-256</td>
<td valign="top" align="center">2.57</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">80.0</td>
</tr> <tr>
<td valign="top" align="left">STSA (Wang Y. et al., <xref ref-type="bibr" rid="B134">2023</xref>)</td>
<td valign="top" align="left">STSFormer-2-256</td>
<td valign="top" align="center">1.99</td>
<td valign="top" align="center">16</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">79.93</td>
</tr>
<tr>
<td valign="top" align="left">SGLFormer (Zhang et al., <xref ref-type="bibr" rid="B161">2024</xref>)</td>
<td valign="top" align="left">SGLFormer-3-256</td>
<td valign="top" align="center">2.58</td>
<td valign="top" align="center">10</td>
<td valign="top" align="left">CIFAR10-DVS</td>
<td valign="top" align="center">82.9</td>
</tr></tbody>
</table>
</table-wrap>
<sec>
<title>4.1 Transformer-based spiking neural networks</title>
<p>Transformer, originally designed for natural language processing (Vaswani et al., <xref ref-type="bibr" rid="B122">2017</xref>), has achieved great success in many computer vision tasks, including image classification (Dosovitskiy et al., <xref ref-type="bibr" rid="B28">2021</xref>; Yuan et al., <xref ref-type="bibr" rid="B159">2021</xref>), object detection (Carion et al., <xref ref-type="bibr" rid="B16">2020</xref>; Liu et al., <xref ref-type="bibr" rid="B82">2021</xref>; Zhu X. et al., <xref ref-type="bibr" rid="B179">2021</xref>), and semantic segmentation (Wang et al., <xref ref-type="bibr" rid="B131">2021</xref>; Yuan et al., <xref ref-type="bibr" rid="B158">2022</xref>). While convolution-based models mainly rely on inductive bias and focus on adjacent pixels, transformer structures use self-attention to capture the relation among spiking features globally, which enhances the performance effectively.</p>
<p>To adopt transformer structure in SNNs, Zhou Z. et al. (<xref ref-type="bibr" rid="B174">2023</xref>) designed a novel spike-form self-attention named Spiking Self Attention (SSA), using sparse spike-form Query, Key and Value without softmax operation. The calculation process of SSA is formulated as <xref ref-type="disp-formula" rid="E18">Equations (18</xref>&#x02013;<xref ref-type="disp-formula" rid="E20">20)</xref>:</p>
<disp-formula id="E18"><label>(18)</label><mml:math id="M20"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>Q</mml:mi><mml:mo>=</mml:mo><mml:mtext>S</mml:mtext><mml:msub><mml:mrow><mml:mtext>N</mml:mtext></mml:mrow><mml:mrow><mml:mi>Q</mml:mi></mml:mrow></mml:msub><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mtext>BN</mml:mtext><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>X</mml:mi><mml:msub><mml:mrow><mml:mi>W</mml:mi></mml:mrow><mml:mrow><mml:mi>Q</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo><mml:mi>K</mml:mi><mml:mo>=</mml:mo><mml:mtext>S</mml:mtext><mml:msub><mml:mrow><mml:mtext>N</mml:mtext></mml:mrow><mml:mrow><mml:mi>K</mml:mi></mml:mrow></mml:msub><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mtext>BN</mml:mtext><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>X</mml:mi><mml:msub><mml:mrow><mml:mi>W</mml:mi></mml:mrow><mml:mrow><mml:mi>K</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr><mml:mtr><mml:mtd><mml:mi>V</mml:mi><mml:mo>=</mml:mo><mml:mtext>S</mml:mtext><mml:msub><mml:mrow><mml:mtext>N</mml:mtext></mml:mrow><mml:mrow><mml:mi>V</mml:mi></mml:mrow></mml:msub><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mtext>BN</mml:mtext><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>X</mml:mi><mml:msub><mml:mrow><mml:mi>W</mml:mi></mml:mrow><mml:mrow><mml:mi>V</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E19"><label>(19)</label><mml:math id="M22"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:msup><mml:mrow><mml:mo class="qopname">SSA</mml:mo></mml:mrow><mml:mrow><mml:mi>&#x02032;</mml:mi></mml:mrow></mml:msup><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>Q</mml:mi><mml:mo>,</mml:mo><mml:mi>K</mml:mi><mml:mo>,</mml:mo><mml:mi>V</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mtext>SN</mml:mtext><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>Q</mml:mi><mml:msup><mml:mrow><mml:mi>K</mml:mi></mml:mrow><mml:mrow><mml:mtext>T</mml:mtext></mml:mrow></mml:msup><mml:mi>V</mml:mi><mml:mo>*</mml:mo><mml:mi>s</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E20"><label>(20)</label><mml:math id="M23"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mo class="qopname">SSA</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>Q</mml:mi><mml:mo>,</mml:mo><mml:mi>K</mml:mi><mml:mo>,</mml:mo><mml:mi>V</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>=</mml:mo><mml:mtext>SN</mml:mtext><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mo class="qopname">BN</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mo class="qopname">Linear</mml:mo><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:msup><mml:mrow><mml:mo class="qopname">SSA</mml:mo></mml:mrow><mml:mrow><mml:mi>&#x02032;</mml:mi></mml:mrow></mml:msup><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>Q</mml:mi><mml:mo>,</mml:mo><mml:mi>K</mml:mi><mml:mo>,</mml:mo><mml:mi>V</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>where <italic>Q, K, V</italic> &#x02208; &#x0211D;<sup><italic>T</italic>&#x000D7;<italic>N</italic>&#x000D7;<italic>D</italic></sup>. The spike-form Query (<italic>Q</italic>), Key (<italic>K</italic>), and Value (<italic>V</italic>) are computed by learnable layers. <italic>s</italic> is a scaling factor, which can be fused into the next spiking neuron in practice. Therefore, the calculation of SSA avoids multiplication, meeting the property of SNNs. Based on the SSA, Zhou Z. et al. (<xref ref-type="bibr" rid="B174">2023</xref>) developed a spiking transformer named Spikformer, which is shown in <xref ref-type="fig" rid="F3">Figure 3</xref>. As the first transformer-based SNN model, Spikformer achieves 74% accuracy on ImageNet-1k, showing great performance potential.</p>
<fig id="F3" position="float">
<label>Figure 3</label>
<caption><p>The overview of spiking transformer (Spikformer).</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnins-18-1383844-g0003.tif"/>
</fig>
<p>Zhou et al. (<xref ref-type="bibr" rid="B169">2023a</xref>) discussed the non-spike computation problem (integer-float multiplications) of Spikformer (Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>) and SEW-ResNet (Fang et al., <xref ref-type="bibr" rid="B35">2021a</xref>), which is caused by Activation-after-addition shortcut. Spikingformer (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>) was proposed with the Pre-activation shortcut to avoid the non-spike computation problem in synaptic computing. Experimental Analysis has shown that Spikingformer has only about 43% energy consumption compared with Spikformer in synaptic computing, with only accumulation operations and lower fire rates. CML (Zhou et al., <xref ref-type="bibr" rid="B171">2023b</xref>) designed a downsampling structure specifically for SNNs to solve the imprecise gradient backpropagation problem of most state-of-the-art deep SNNs (including Spikformer). CML achieved 77.34% on ImageNet, significantly enhancing the performance of transformer-based SNNs. All the architectures above are based on SSA with computational complexity of <italic>O</italic>(<italic>N</italic><sup>2</sup><italic>d</italic>) or <italic>O</italic>(<italic>Nd</italic><sup>2</sup>), while Yao et al. (<xref ref-type="bibr" rid="B152">2023a</xref>) designed a novel Spike-Driven Self-Attention (SDSA) with linear complexity regarding both the number of tokens and channels. SDSA uses only mask and addition operations without any multiplication, thus having up to 87.2 &#x000D7; lower computation energy than the vanilla SSA. In addition, the Spike-driven Transformer based on SDSA has achieved 77.1% accuracy on ImageNet-1k. Wang Y. et al. (<xref ref-type="bibr" rid="B134">2023</xref>) proposed an SNN-based spatial-temporal self-attention (STSA) mechanism, which could calculate the feature dependence across the time and space domains. Shi et al. (<xref ref-type="bibr" rid="B118">2024</xref>) proposed Dual Spike Self-Attention (DSSA) with a reasonable scaling method, achieving 79.40% top-1 accuracy on ImageNet-1K. Yao et al. (<xref ref-type="bibr" rid="B149">2024</xref>) proposed Spike-driven Transformer v2 which explored the impact of structure, spike-driven self-attention, and skip connection on its performance to inspire the next-generation transformer-based neuromorphic chip designs. Zhou Z. et al. (<xref ref-type="bibr" rid="B173">2024</xref>) developed a Spiking Convolutional Stem (SCS) with supplementary layers to enhance the architecture of Spikformer, achieving 80.38% accuracy on ImageNet-1k. Zhang et al. (<xref ref-type="bibr" rid="B161">2024</xref>) proposed a Spiking Global-Local-Fusion Transformer (SGLFormer), which enables efficient information processing on both global and local scales, by integrating transformer and convolution structures in SNNs. SGLFormer achieved a groundbreaking top-1 accuracy of 83.73% on ImageNet-1k with 64M parameters. Zhou C. et al. (<xref ref-type="bibr" rid="B170">2024</xref>) proposed QKFormer, a novel hierarchical spiking transformer using Q-K attention, which can easily model the importance of token or channel dimensions with binary values and has linear complexity to &#x00023;tokens (or &#x00023;channels). QKFormer achieved a significant milestone, surpassing 85% top-1 accuracy on ImageNet with 4 time steps using the direct training approach.</p>
<p>Biological realistic models tend to model neural networks with high biological plausibility to simulate the complex biological mechanism of the brain. It often lacks the consideration of computational efficiency and performance optimization on general application tasks. Traditional ANNs often prioritize task performance over biological realism and computational energy consumption. SNNs have great potential to own the characteristics of biological plausibility, low computational energy consumption, and high task performance simultaneously. Especially, several direct training Transformer-based SNNs have broken through 80% top-1 accuracy on ImageNet-1K, which instills great optimism in the application of SNNs.</p>
</sec>
<sec>
<title>4.2 Residual architectures in spiking neural networks</title>
<p>Residual block is the fundamental block in both deep ANNs and SNNs. As shown in <xref ref-type="fig" rid="F4">Figure 4</xref>, there are mainly three residual shortcut types in SNNs: Activation-after-addition, Activation-before-addition, and Pre-activation. Both advantages and disadvantages of these three types are concluded in <xref ref-type="table" rid="T3">Table 3</xref>. <bold>Activation-after-addition shortcut</bold> simply replaces ReLU activation layers in the standard residual block with spiking neurons, such as Spiking ResNet (Hu et al., <xref ref-type="bibr" rid="B50">2021a</xref>) and MPBN (Guo et al., <xref ref-type="bibr" rid="B45">2023b</xref>). SNNs with this simple design suffer from performance degradation and gradient vanishing/exploding. For example, the deeper 34-layer Spiking ResNet has lower test accuracy than the shallower 18-layer Spiking ResNet. As the layer increases, the test accuracy of Spiking ResNet decreases (Fang et al., <xref ref-type="bibr" rid="B35">2021a</xref>). To solve the degradation problem in the Activation-after-addition shortcut, <bold>Activation-before-addition shortcut</bold> is proposed in SEW-ResNet (Fang et al., <xref ref-type="bibr" rid="B35">2021a</xref>), which extended directly trained SNNs to 100 layers for the first time. This structure has been widely used, such as in Spikformer (Zhou Z. et al., <xref ref-type="bibr" rid="B174">2023</xref>), PLIF (Fang et al., <xref ref-type="bibr" rid="B33">2021b</xref>), PSN (Fang et al., <xref ref-type="bibr" rid="B34">2023b</xref>). This design mitigates the vanishing/exploding gradient problem and could train deeper SNN. However, the blocks in this shortcut will result in positive integers, which leads to non-spike computations (integer-float multiplications) in synaptic computing (like convolutional layer, linear layer) (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>). <bold>Pre-activation shortcut</bold> could be traced back to the <italic>Activation</italic>-<italic>Conv</italic>-<italic>Bn</italic> paradigm, which is a fundamental building block in Binary Neural Networks (BNNs) (Liu et al., <xref ref-type="bibr" rid="B81">2018</xref>, <xref ref-type="bibr" rid="B80">2020</xref>; Guo N. et al., <xref ref-type="bibr" rid="B41">2021</xref>; Zhang Y. et al., <xref ref-type="bibr" rid="B165">2022</xref>). Some representative SNNs that use the Pre-activation shortcut include MS-ResNet (Hu et al., <xref ref-type="bibr" rid="B51">2021b</xref>), Spikingformer (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>), Spike-driven transformer (Yao et al., <xref ref-type="bibr" rid="B152">2023a</xref>). MS-ResNet directly trained convolution-based SNNs to successfully extend the depth up to 482 layers on CIFAR10 without experiencing degradation problems, effectively verifying the feasibility of this way. Spikingformer (Zhou et al., <xref ref-type="bibr" rid="B169">2023a</xref>) showed that the Pre-activation shortcut can effectively avoid non-spike computations, and thus has lower energy consumption than the previous shortcut in synaptic computing, through avoiding integer-float multiplication problems and with a lower firing rate. However, the Pre-activation shortcut requires dense transmission of floats in the residual branch.</p>
<fig id="F4" position="float">
<label>Figure 4</label>
<caption><p>The overview of residual learning architectures.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnins-18-1383844-g0004.tif"/>
</fig>
<table-wrap position="float" id="T3">
<label>Table 3</label>
<caption><p>Features of various residual learning architectures.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>Features</bold></th>
<th valign="top" align="left"><bold>Activation-after-addition</bold></th>
<th valign="top" align="left"><bold>Activation-before-addition</bold></th>
<th valign="top" align="left"><bold>Pre-activation</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Element addition</td>
<td valign="top" align="left">Spike added to floats</td>
<td valign="top" align="left">Integer added to spike</td>
<td valign="top" align="left">Floats added to floats</td>
</tr> <tr>
<td valign="top" align="left">Gradient vanishing/exploding</td>
<td valign="top" align="left">Yes</td>
<td valign="top" align="left">No</td>
<td valign="top" align="left">No</td>
</tr> <tr>
<td valign="top" align="left">Synaptic computing type</td>
<td valign="top" align="left">Spike computing</td>
<td valign="top" align="left">Multiplication between sparse integer and floats</td>
<td valign="top" align="left">Spike computing</td>
</tr>
<tr>
<td valign="top" align="left">Data transmission in residual branch</td>
<td valign="top" align="left">Spike</td>
<td valign="top" align="left">Sparse integer</td>
<td valign="top" align="left">Floats</td>
</tr></tbody>
</table>
</table-wrap>
<p>Overall, the residual learning suitable for the properties of SNNs needs further exploration. In our opinion, the Activation-after-addition shortcut with gradient problem is not suitable for directly training deep SNNs, but is feasible in the field of ANN-to-SNN conversion. Activation-before-addition shortcut has some alternatives to ensure the properties of SNNs by slightly sacrificing the performance, such as using AND or IAND to replace ADD in the aggregation operation. Pre-activation shortcut needs further analyses of the effects of float transmission, and more efforts to exploit its advantages through collaborative hardware optimization and design.</p>
</sec>
<sec>
<title>4.3 Others</title>
<p>Besides the above-mentioned architectures, some other interesting research topics are also worthy of attention, such as Spiking RNN/LSTM, LSM, etc. Deep Liquid State Machine (LSM) (Wang and Li, <xref ref-type="bibr" rid="B127">2016</xref>) explored the power of recurrent spiking networks and deep architectures. Soures and Kudithipudi (<xref ref-type="bibr" rid="B120">2019</xref>) proposed a novel deep LSM to capture dynamic information over multiple time-scales with a combination of randomly connected layers and unsupervised layers. Hamilton et al. (<xref ref-type="bibr" rid="B47">2019</xref>) demonstrated the nonlinear dynamics of spiking neurons can be used to implement low-level graph operations. Zhu Z. et al. (<xref ref-type="bibr" rid="B180">2022</xref>) proposed end-to-end Spiking Graph Convolutional Networks (GCNs) that integrate the embedding of GCNs with the biofidelity characteristics of SNNs. Bellec et al. (<xref ref-type="bibr" rid="B8">2020</xref>) and Bohnstingl et al. (<xref ref-type="bibr" rid="B12">2022</xref>) explored the architectures and online-training methods of recurrent spiking neural networks. Ren H. et al. (<xref ref-type="bibr" rid="B110">2023</xref>) proposed a novel end-to-end point-based SNN architecture, which excels at processing sparse event cloud data, effectively extracting both global and local features through a singular-stage structure.</p>
</sec>
</sec>
<sec id="s5">
<title>5 Software frameworks and neuromorphic hardware for spiking neural networks</title>
<sec>
<title>5.1 Software frameworks for training spiking neural networks</title>
<p>Software frameworks play a crucial role in propelling the advancement of deep learning. Deep learning frameworks such as PyTorch (Paszke et al., <xref ref-type="bibr" rid="B99">2019</xref>) and TensorFlow (Abadi et al., <xref ref-type="bibr" rid="B1">2016</xref>) leverage low-level languages like C&#x0002B;&#x0002B; libraries for high-performance acceleration on the backend, while offering user-friendly front-end application programming interfaces (APIs) implemented in high-level languages like Python. These frameworks significantly ease the workload of constructing and training ANNs, making substantial contributions to the growth of deep learning research. However, these deep learning frameworks are primarily designed for ANNs. With the development of large-scale brain-inspired neural networks, many related frameworks have emerged, facilitating the modeling and efficient computation of large-scale SNNs.</p>
<p>One category of frameworks includes brain simulators such as NEURON (Hines and Carnevale, <xref ref-type="bibr" rid="B49">1997</xref>) and Brian (Goodman and Brette, <xref ref-type="bibr" rid="B40">2009</xref>), which not only enhance the scalability and computational efficiency of models but also encompass cognitive functions such as perception, decision-making, and reasoning. The SNNs constructed by these frameworks exhibit a high degree of biological plausibility, making them suitable for studying the functionalities of real neural systems. They support biologically interpretable learning rules such as Spike-Timing-Dependent Plasticity (STDP) (Bi and Poo, <xref ref-type="bibr" rid="B10">1998</xref>), playing a significant role in advancing the field of neuroscience. However, these frameworks lack core computational functionalities required for deep learning, such as automatic differentiation, rendering them incapable of performing machine learning tasks.</p>
<p>Another category of brain-inspired computing frameworks comprises deep spiking computation frameworks. Deep SNNs involve a substantial amount of matrix operations across spatial and temporal dimensions, a variety of neurons, neuromorphic datasets, and deployments on neuromorphic chips. The modeling and application processes are complex, and achieving high-performance acceleration is challenging. To address these issues, spiking deep learning frameworks need to support the construction, training, and deployment of deep SNNs, and be capable of acceleration based on spike operations. Frameworks such as BindsNET (Hazan et al., <xref ref-type="bibr" rid="B48">2018</xref>), NengoDL (Rasmussen, <xref ref-type="bibr" rid="B106">2019</xref>), SpykeTorch (Mozafari et al., <xref ref-type="bibr" rid="B93">2019</xref>), <ext-link ext-link-type="uri" xlink:href="https://doi.org/10.5281/zenodo.4422025">Norse</ext-link>, <ext-link ext-link-type="uri" xlink:href="https://github.com/fzenke/spytorch">SpyTorch</ext-link>, SNNTorch (Eshraghian et al., <xref ref-type="bibr" rid="B31">2023</xref>), and SpikingJelly (Fang et al., <xref ref-type="bibr" rid="B32">2023a</xref>) have been developed. They utilize simple spiking neurons to reduce computational complexity, making them suitable for machine learning research. Among them, BindsNet (Hazan et al., <xref ref-type="bibr" rid="B48">2018</xref>) primarily focuses on machine learning and reinforcement learning; NengoDL (Rasmussen, <xref ref-type="bibr" rid="B106">2019</xref>) converts ANNs to obtain deep SNNs but does not support direct training of SNNs using surrogate gradient methods; SpyTorch is a demonstrative framework that only provides basic surrogate gradient examples; SpyTorch (Mozafari et al., <xref ref-type="bibr" rid="B93">2019</xref>) introduces a new type of surrogate gradient method named SuperSpike. These frameworks can implement some simple machine learning and reinforcement learning models, but they still lack deep learning capabilities for SNNs. Norse is attempting to introduce the sparse and event-driven characteristics of SNNs and supports many typical spiking neuron models. It is in the development stage and has not been officially released yet. SNNTorch supports some variants of online backpropagation algorithms that are more biologically plausible and support large-scale SNN computation. SpikingJelly (Fang et al., <xref ref-type="bibr" rid="B32">2023a</xref>) is a full-stack toolkit for preprocessing neuromorphic datasets, building deep SNNs, optimizing their parameters, and deploying SNNs on neuromorphic chips, which shows remarkable extensibility and flexibility, enabling users to accelerate custom models at low costs through multilevel inheritance and semiautomatic code generation. In summary, the development of existing software frameworks is essentially in its early stages, and there is still a long way to go in terms of functionality enhancement and performance optimization.</p>
</sec>
<sec>
<title>5.2 Neuromorphic hardware for spiking neural networks</title>
<p>Neuromorphic hardware provides computational power for neural network models, playing a crucial role in large-scale brain-like neural networks. Efficient hardware can significantly accelerate the training, evaluation, iteration, and real-world applications of large-scale brain-like models. In comparison to general-purpose processors, deep learning chips and brain-like chips are specialized chips that focus on the computational efficiency of deep learning tasks and brain-like computing tasks, aiming to achieve better power/performance/area ratios. Current deep learning chips, like general CPUs, are based on the Von Neumann architecture, with separate computing and storage units. Brain-like chips enhance computational efficiency by designing efficient storage and computation hierarchy, enabling parallel data flow and efficient reuse, thus improving computational efficiency. From an architectural perspective, current brain-like chips can be mainly divided into two categories: analog-digital hybrid circuits and fully digital circuits (<xref ref-type="table" rid="T4">Table 4</xref>).</p>
<table-wrap position="float" id="T4">
<label>Table 4</label>
<caption><p>Overview of typical neuromorphic hardware.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>Chip</bold></th>
<th valign="top" align="left"><bold>Developer</bold></th>
<th valign="top" align="left"><bold>Network</bold></th>
<th valign="top" align="left"><bold>Function</bold></th>
<th valign="top" align="left"><bold>Arch</bold></th>
<th valign="top" align="left"><bold>Scale</bold></th>
</tr>
</thead>
<tbody>
<tr style="background-color:#dee1e1">
<td valign="top" align="left" colspan="6"><bold>Hybrid digital-analog</bold></td>
</tr> <tr>
<td valign="top" align="left">BrainScaleS (Schemmel et al., <xref ref-type="bibr" rid="B113">2010</xref>)</td>
<td valign="top" align="left">Heidelberg Uni</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr>
<tr>
<td valign="top" align="left">Neurogrid (Benjamin et al., <xref ref-type="bibr" rid="B9">2014</xref>)</td>
<td valign="top" align="left">Stanford</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Inference</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">ROLLS (Qiao et al., <xref ref-type="bibr" rid="B103">2015</xref>)</td>
<td valign="top" align="left">UZH</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Small</td>
</tr> <tr>
<td valign="top" align="left">DYNAPs (Moradi et al., <xref ref-type="bibr" rid="B92">2017</xref>)</td>
<td valign="top" align="left">UZH</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Inference</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Small</td>
</tr> <tr>
<td valign="top" align="left">Memristor-based (Zhang et al., <xref ref-type="bibr" rid="B164">2021</xref>)</td>
<td valign="top" align="left">&#x02013;</td>
<td valign="top" align="left">ANNs/SNNs</td>
<td valign="top" align="left">Inference</td>
<td valign="top" align="left">CIM</td>
<td valign="top" align="left">Small</td>
</tr> <tr>
<td valign="top" align="left">BrainScaleS-2 (Pehle et al., <xref ref-type="bibr" rid="B100">2022</xref>)</td>
<td valign="top" align="left">Heidelberg Uni</td>
<td valign="top" align="left">ANNs/SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr style="background-color:#dee1e1">
<td valign="top" align="left" colspan="6"><bold>Digital</bold></td>
</tr> <tr>
<td valign="top" align="left">SpiNNaker (Painkras et al., <xref ref-type="bibr" rid="B95">2013</xref>)</td>
<td valign="top" align="left">UoM</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">SpiNNaker 2</td>
<td valign="top" align="left">UoM</td>
<td valign="top" align="left">ANNs/SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">TrueNorth (Akopyan et al., <xref ref-type="bibr" rid="B3">2015</xref>)</td>
<td valign="top" align="left">IBM</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Inference</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">Darwin (Shen et al., <xref ref-type="bibr" rid="B117">2016</xref>)</td>
<td valign="top" align="left">ZJU</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Inference</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Small</td>
</tr> <tr>
<td valign="top" align="left">Darwin II (Ma et al., <xref ref-type="bibr" rid="B86">2017</xref>)</td>
<td valign="top" align="left">ZJU</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">Darwin III (Ma et al., <xref ref-type="bibr" rid="B85">2024</xref>)</td>
<td valign="top" align="left">ZJU</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">DeepSouth (Wang et al., <xref ref-type="bibr" rid="B128">2017</xref>)</td>
<td valign="top" align="left">Westwell</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Inference</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">Intel SNN chip (Chen et al., <xref ref-type="bibr" rid="B21">2018</xref>)</td>
<td valign="top" align="left">Intel</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">ODIN (Frenkel et al., <xref ref-type="bibr" rid="B38">2018</xref>)</td>
<td valign="top" align="left">K.U.Leuven</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Small</td>
</tr> <tr>
<td valign="top" align="left">Loihi (Davies et al., <xref ref-type="bibr" rid="B26">2018</xref>)</td>
<td valign="top" align="left">Intel</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">Tianjic (Pei et al., <xref ref-type="bibr" rid="B101">2019</xref>)</td>
<td valign="top" align="left">Tsinghua</td>
<td valign="top" align="left">ANNs/SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">MorphIC (Frenkel et al., <xref ref-type="bibr" rid="B39">2019</xref>)</td>
<td valign="top" align="left">UZH</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Small</td>
</tr> <tr>
<td valign="top" align="left">Flash-based (Wu et al., <xref ref-type="bibr" rid="B138">2020</xref>)</td>
<td valign="top" align="left">&#x02013;</td>
<td valign="top" align="left">ANNs/SNNs</td>
<td valign="top" align="left">Inference</td>
<td valign="top" align="left">CIM</td>
<td valign="top" align="left">Small</td>
</tr> <tr>
<td valign="top" align="left">Loihi II (Davies, <xref ref-type="bibr" rid="B25">2021</xref>)</td>
<td valign="top" align="left">Intel</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">Y. Kuang et al. (Kuang et al., <xref ref-type="bibr" rid="B62">2021</xref>)</td>
<td valign="top" align="left">PKU</td>
<td valign="top" align="left">ANNs/SNNs</td>
<td valign="top" align="left">Inference</td>
<td valign="top" align="left">NMA</td>
<td valign="top" align="left">Large</td>
</tr> <tr>
<td valign="top" align="left">H2Learn (Liang et al., <xref ref-type="bibr" rid="B77">2021</xref>)</td>
<td valign="top" align="left">UCSB</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">SNN TA</td>
<td valign="top" align="left">Large</td>
</tr>
<tr>
<td valign="top" align="left">SATA (Yin et al., <xref ref-type="bibr" rid="B155">2022</xref>)</td>
<td valign="top" align="left">Yale</td>
<td valign="top" align="left">SNNs</td>
<td valign="top" align="left">Training</td>
<td valign="top" align="left">SNN TA</td>
<td valign="top" align="left">Small</td>
</tr></tbody>
</table>
<table-wrap-foot>
<p>SNN TA, SNN training accelerator; Arch, architecture.</p>
</table-wrap-foot>
</table-wrap>
<p>Inspired by the simultaneous computation and storage capabilities of the brain&#x00027;s neural system, brain-inspired chips often adopt near-memory (NMA) or compute-in-memory architectures (CIM), incorporating closely coupled computational and storage resources within each computing core (Akopyan et al., <xref ref-type="bibr" rid="B3">2015</xref>; Pei et al., <xref ref-type="bibr" rid="B101">2019</xref>). Efficient intra-chip and inter-chip interconnects enable large-scale computational parallelism and high local memory, reducing computational power consumption.</p>
<p>The near-memory computing architecture refers to the separation of memory storage and computation in each processing unit, but with proximity. Key chips in this category include IBM&#x00027;s TrueNorth (Akopyan et al., <xref ref-type="bibr" rid="B3">2015</xref>), Intel&#x00027;s Loihi (Davies et al., <xref ref-type="bibr" rid="B26">2018</xref>; Davies, <xref ref-type="bibr" rid="B25">2021</xref>), the University of Manchester&#x00027;s SpiNNaker (Painkras et al., <xref ref-type="bibr" rid="B95">2013</xref>), Stanford University&#x00027;s Neurogrid (Benjamin et al., <xref ref-type="bibr" rid="B9">2014</xref>), Heidelberg University&#x00027;s BrainScaleS (Schemmel et al., <xref ref-type="bibr" rid="B113">2010</xref>; Pehle et al., <xref ref-type="bibr" rid="B100">2022</xref>), Tsinghua University&#x00027;s Tianji Chip (Pei et al., <xref ref-type="bibr" rid="B101">2019</xref>), and Zhejiang University&#x00027;s Darwin Chip (Shen et al., <xref ref-type="bibr" rid="B117">2016</xref>; Ma et al., <xref ref-type="bibr" rid="B86">2017</xref>, <xref ref-type="bibr" rid="B85">2024</xref>). They utilize characteristics of brain-like spiking computation such as sparsity, spike summation, and asynchronous event-driven processing to achieve ultra-low power consumption, currently mainly supporting model inference and local online learning based on STDP, Hebb, etc. For instance, ROLLS (Qiao et al., <xref ref-type="bibr" rid="B103">2015</xref>), ODIN (Frenkel et al., <xref ref-type="bibr" rid="B38">2018</xref>), and MorphIC (Frenkel et al., <xref ref-type="bibr" rid="B39">2019</xref>) support spike-driven synaptic plasticity (SDSP) rules, and Loihi adds a learning module for STDP rules. In SpiNNaker (Painkras et al., <xref ref-type="bibr" rid="B95">2013</xref>) and BrainScaleS (Schemmel et al., <xref ref-type="bibr" rid="B113">2010</xref>), STDP learning is exhibited through timestamp recording and learning circuits. In their next generations (Pehle et al., <xref ref-type="bibr" rid="B100">2022</xref>), more flexible learning rules are possible due to the presence of embedded programmable units. Tsinghua University&#x00027;s Tianji Chip, as the first chip to support the fusion of SNN and ANN computation, improves accuracy based on ANN, and achieves rich dynamics, high efficiency, and robustness based on SNN. This mode is also adopted by BrainScaleS-2 (Pehle et al., <xref ref-type="bibr" rid="B100">2022</xref>), SpiNNaker-2, and Loihi-2 (Davies, <xref ref-type="bibr" rid="B25">2021</xref>). Recently, BPTT has been applied to SNNs, achieving higher accuracy compared to local learning rules (Wu et al., <xref ref-type="bibr" rid="B140">2018</xref>, <xref ref-type="bibr" rid="B141">2019</xref>). Some works, like H2Learn (Liang et al., <xref ref-type="bibr" rid="B77">2021</xref>) and SATA (Yin et al., <xref ref-type="bibr" rid="B155">2022</xref>), have designed specific architectures for BPTT learning in SNNs. In the future, the integration of learning rules will become increasingly important for exploring large and complex neuromorphic models in brain-inspired computing (BIC) chips.</p>
<p>Another important type of BIC architecture is the compute-in-memory architecture, where in-core processing units and on-chip storage are physically integrated, performing synaptic integration matrix operations in synaptic memory. Compute-in-memory chips can be divided into two categories based on the materials: traditional or emerging memories. Traditional memories (such as SRAM, DRAM, and Flash) can be redesigned to support specific logical operations (Wu et al., <xref ref-type="bibr" rid="B138">2020</xref>). Their advantages include a mature ecosystem, easy simulation, and manufacturing. Emerging memories mainly refer to storage devices based on memristors. Synaptic weight storage, multiplication calculations, and presynaptic inputs are performed at the same crosspoint in the memristor, integrating computation and storage. Brain-inspired computing hardware based on memristors involves multiple levels of material and architectural designs, which is currently still in a small-scale phase due to manufacturing process limitations.</p>
<p>Multiple types of brain-like chips have shown remarkable developments, demonstrating significant advantages in terms of biological simulation and low-power inference. However, they still face numerous challenges in practical applications. When it comes to handling high-level intelligence tasks, the superiority of brain-like chips compared to GPUs and ANN accelerators has not been fully established. Currently, to optimize their performance, some designs draw inspiration from ANN accelerators for improvements. It&#x00027;s worth noting that current brain-like chips do not yet support the training of large-scale SNNs and require special architectural designs to accommodate the training process for SNNs. To further support large-scale SNNs, it is necessary to enhance brain-like systems from both a software and hardware perspective in a more collaborative manner.</p>
</sec>
<sec>
<title>5.3 Software and hardware interplay</title>
<p>The deployment of algorithms for SNNs onto neuromorphic chips typically requires certain software frameworks. The computational software frameworks mentioned in Section 5.1 usually support simulations on mainstream CPUs and GPUs, without clear mention of deployment on neuromorphic chips. Meanwhile, among the previously mentioned neuromorphic chips in Section 5.2, only about 27% of them are connected to application software packages (Schuman et al., <xref ref-type="bibr" rid="B114">2022</xref>). Typically, these application software packages contain model construction tools, simulators, and optimization tools. Model construction tools are used to define the structure and parameters of neural networks, including neuron types, connection patterns. Simulators are applied to simulate and debug neural network models on the chip. Optimization tools are adopted to train the network parameters and optimize its performance. Here are some typical examples. The Neurogrid chip is paired with the Neurogrid Software Framework (Benjamin et al., <xref ref-type="bibr" rid="B9">2014</xref>), allowing users to specify neuronal models in the Python programming environment. The software framework for the BrainScaleS chip is the BrainScaleS-Software-Stack (Pehle et al., <xref ref-type="bibr" rid="B100">2022</xref>), which supports training neural networks on the chip using the PyTorch framework. The IBM TrueNorth chip typically utilizes a software framework called the TrueNorth Ecosystem, which is developed by the TrueNorth native Corelet language (Akopyan et al., <xref ref-type="bibr" rid="B3">2015</xref>). IBM NorthPole (Modha et al., <xref ref-type="bibr" rid="B91">2023</xref>) is a brain-inspired memory-near-compute chip with a software development kit, but this chip does not emulate spiking communication. Tianjic&#x00027;s software toolchain supports both ANN-to-SNN conversion and direct training for SNNs, and supports automatically transforming a pretrained model into an equivalent network that meets the Tianjic hardware constraints for non-spiking ANNs (Pei et al., <xref ref-type="bibr" rid="B101">2019</xref>). The latest Darwin3 builds a specialized instruction set architecture (ISA) (Ma et al., <xref ref-type="bibr" rid="B85">2024</xref>), which is close to machine code tailored for efficient neuromorphic computing. These software frameworks enable users to conveniently construct, simulate, and optimize neural network models on neuromorphic chips, facilitating efficient research and application development.</p>
</sec>
</sec>
<sec id="s6">
<title>6 Applications of deep spiking neural networks</title>
<p>SNNs offer powerful computation capability due to their event-driven nature and temporal processing property. Theoretically, SNNs could be applied to any field where conventional deep neural networks (DNNs) are applied. As the training methods and programming frameworks of deep SNNs become more powerful, deep SNNs are increasingly drawing more attention and being applied to more fields, mainly including computer vision, reinforcement learning and autonomous robotics, biological visual system modeling, biological signal processing, natural language processing, equipment safety monitoring, and so on. It should be noted that this paper only lists some typical examples in recent years for some common application fields, not aiming to fully review all related studies.</p>
<sec>
<title>6.1 Applications in computer vision</title>
<p>As traditional DNNs, the most common applications of SNNs lay in computer vision tasks. There are mainly two types of visual inputs for SNNs, i.e., RGB frames from traditional cameras or events from neuromorphic vision sensors. Neuromorphic vision sensors display great potential for computer vision tasks under high-speed and low-light conditions (Li and Tian, <xref ref-type="bibr" rid="B68">2021</xref>). SNNs are excellent candidates for processing neuromorphic signals due to their event-driven nature and energy-efficient computing.</p>
<p>Recognition task plays an important role in the rapid progress of deep SNNs. SNNs are usually tested on both static datasets such as CIFAR10, CIFAR100, ImageNet, and neuromorphic datasets such as CIFAR10-DVS and DVS128-Gesture. <xref ref-type="table" rid="T2">Table 2</xref> lists the performances of some recently proposed architectures. Besides common recognition tasks, deep SNNs are increasingly applied to more computer vision tasks, including object detection/tracking, image denoising/generation, image/video reconstruction, video action recognition, image segmentation, and so on.</p>
<sec>
<title>6.1.1 Object detection and object tracking</title>
<p>The first spike-based object detection model Spiking-YOLO was obtained through the ANN-to-SNN conversion method, achieving comparable performances to tiny-YOLO on PASCAL VOC and MS-COCO dataset with 3,500 time steps (Kim et al., <xref ref-type="bibr" rid="B57">2020</xref>). Later, a spike calibration (SpiCalib) method was proposed to reduce the time steps to hundreds (Li et al., <xref ref-type="bibr" rid="B74">2022</xref>). Kugele et al. (<xref ref-type="bibr" rid="B63">2021</xref>) and Cordone et al. (<xref ref-type="bibr" rid="B23">2022</xref>) combined some spiking backbones with an SSD detection head for event cameras.</p>
<p>Considering that the Siamese networks have achieved remarkable performances in object tracking, SiamSNN was constructed by conversion to achieve short latency and low precision degradation on several benchmarks (Luo et al., <xref ref-type="bibr" rid="B83">2022</xref>). Similarly, the directly trained Spiking SiamFC&#x0002B;&#x0002B; showed a small precision loss compared to the original SiamFC&#x0002B;&#x0002B; (Xiang et al., <xref ref-type="bibr" rid="B142">2022</xref>). A spiking transformer network called STNet was developed for event-based single-object tracking, demonstrating competitive tracking accuracy and speed on three event-based datasets (Zhang et al., <xref ref-type="bibr" rid="B167">2022a</xref>). To process frames and events simultaneously, Yang et al. (<xref ref-type="bibr" rid="B148">2019</xref>) proposed DashNet, achieving good tracking performance with a surprising tracking speed of 2,083 FPS on neuromorphic chips.</p>
</sec>
<sec>
<title>6.1.2 Image generation/denoising and image/video reconstruction</title>
<p>Generation tasks are increasingly explored in SNNs. Com&#x0015F;a et al. (<xref ref-type="bibr" rid="B22">2021</xref>) introduced a directly trained spiking autoencoder to reconstruct images with high fidelity on MNIST and FMNIST datasets. Kamata et al. (<xref ref-type="bibr" rid="B56">2022</xref>) constructed a fully spiking variational autoencoder (FSVAE), generating images with competitive quality compared to conventional ANNs. Liu et al. (<xref ref-type="bibr" rid="B79">2023</xref>) proposed a Spiking-Diffusion model, outperforming the existing SNN-based generation models on several datasets. Castagnetti et al. (<xref ref-type="bibr" rid="B17">2023</xref>) developed an image denoising solution based on a directly trained spiking autoencoder, achieving a competitive signal-to-noise ratio on the Set12 dataset with significantly lower energy.</p>
<p>Visual information reconstruction is important for neuromorphic vision sensors, because humans cannot directly perceive visual scenes from events. Zhu and Tian (<xref ref-type="bibr" rid="B176">2023</xref>) provided a comprehensive review of visual reconstruction methods for events. Duwek et al. (<xref ref-type="bibr" rid="B30">2021</xref>) proposed a hybrid ANN-SNN model, accomplishing image reconstruction for simple scenes from N-MNIST and N-Caltech101 datasets. Zhu L. et al. (<xref ref-type="bibr" rid="B175">2021</xref>) proposed an image reconstruction algorithm that combines DVS and Vidar signals, leveraging the high dynamic range of DVS to improve reconstruction effectiveness. Subsequently, they developed a deep SNN with an encoder-decoder structure for event-based video reconstruction, achieving performance comparable to ANN counterparts with only 0.05x energy consumption (Zhu L. et al., <xref ref-type="bibr" rid="B177">2022</xref>).</p>
</sec>
<sec>
<title>6.1.3 Others</title>
<p>Besides the above-mentioned tasks, deep SNNs are also applied in some other computer vision tasks, including video action recognition (Panda and Srinivasa, <xref ref-type="bibr" rid="B96">2018</xref>; Wang et al., <xref ref-type="bibr" rid="B130">2019</xref>; Zhang et al., <xref ref-type="bibr" rid="B163">2022c</xref>; Chakraborty and Mukhopadhyay, <xref ref-type="bibr" rid="B18">2023</xref>; Yu et al., <xref ref-type="bibr" rid="B157">2024</xref>), image segmentation (Parameshwara et al., <xref ref-type="bibr" rid="B97">2021</xref>; Kim et al., <xref ref-type="bibr" rid="B58">2022</xref>; Liang et al., <xref ref-type="bibr" rid="B76">2022</xref>; Zhang H. et al., <xref ref-type="bibr" rid="B160">2023</xref>), optical flow estimation (Lee et al., <xref ref-type="bibr" rid="B64">2020a</xref>; Cuadrado et al., <xref ref-type="bibr" rid="B24">2023</xref>; Kosta and Roy, <xref ref-type="bibr" rid="B60">2023</xref>), depth prediction (Ran&#x000E7;on et al., <xref ref-type="bibr" rid="B105">2022</xref>; Wu et al., <xref ref-type="bibr" rid="B139">2022</xref>; Zhang et al., <xref ref-type="bibr" rid="B162">2022b</xref>), point clouds processing (Zhou et al., <xref ref-type="bibr" rid="B172">2020</xref>; Ren D. et al., <xref ref-type="bibr" rid="B109">2023</xref>), human pose tracking (Zou et al., <xref ref-type="bibr" rid="B181">2023</xref>), lip-reading (Bulzomi et al., <xref ref-type="bibr" rid="B14">2023</xref>), emotion/expression recognition (Wang B. et al., <xref ref-type="bibr" rid="B123">2022</xref>; Barchid et al., <xref ref-type="bibr" rid="B6">2023</xref>), medical image classification (Shan et al., <xref ref-type="bibr" rid="B115">2022</xref>; Qasim Gilani et al., <xref ref-type="bibr" rid="B102">2023</xref>), and so on.</p>
</sec>
</sec>
<sec>
<title>6.2 Applications in other fields</title>
<p>Besides computer vision tasks, SNNs are showing gradually expanding application prospects in many fields, including reinforcement learning and autonomous robotics, biological visual system modeling, biological signal processing, natural language processing, equipment safety monitoring, and so on.</p>
<sec>
<title>6.2.1 Reinforcement learning and autonomous robotics</title>
<p>As reinforcement learning (RL) is critical for the survival of humans and animals, there is increasing interest in applying brain-inspired SNNs to reinforcement learning. To reduce the latency of spiking RL, Qin et al. (<xref ref-type="bibr" rid="B104">2023</xref>) applied learnable matrix multiplication to encode and decode spikes.</p>
<p>Due to the good biological plausibility and high energy efficiency, SNNs have been applied to autonomous robotics for a long time, which is still a flourishing research direction, mainly including pattern generation (walk, trot, or run), motor control, and navigation (simultaneous localization and mapping, SLAM). Yamazaki et al. (<xref ref-type="bibr" rid="B147">2022</xref>) have already provided a good review of relevant studies, we do not go into more detail about this topic in this review which mainly focuses on deep SNNs.</p>
</sec>
<sec>
<title>6.2.2 Biological visual system modeling and biological signal processing</title>
<p>ANNs play important roles in modeling biological visual pathways. However, SNNs are more biologically plausible models due to the use of temporal spike sequences. Therefore, several studies adopted SNNs to model the biological visual cortex. Further, they added a brain-inspired recurrent module into deep SNNs, outperforming the forward deep SNNs under natural movie stimuli (Huang et al., <xref ref-type="bibr" rid="B52">2023b</xref>). Zhang J. et al. (<xref ref-type="bibr" rid="B166">2023</xref>) compared performances of deep SNNs and CNNs in the prediction of visual responses to naturalistic stimuli in three brain areas. Ma et al. (<xref ref-type="bibr" rid="B87">2023</xref>) presented a temporal conditioning spiking latent variable model to produce more realistic spike activities.</p>
<p>Due to the intrinsic dynamics, SNNs are also applied to process biological signals. Xiong et al. (<xref ref-type="bibr" rid="B145">2021</xref>) proposed a convolutional SNN for odor recognition of electronic noses.</p>
</sec>
<sec>
<title>6.2.3 Others</title>
<p>SNNs were also applied to natural language processing, equipment safety monitoring, semantic communication, multi-modal information processing, and so on. To ease the heavy energy cost of ANN-based large language models, some studies applied SNN-based architectures, including SpikBERT (Lv et al., <xref ref-type="bibr" rid="B84">2023</xref>), SpikingBERT (Bal and Sengupta, <xref ref-type="bibr" rid="B5">2024</xref>), SpikeGPT (Zhu et al., <xref ref-type="bibr" rid="B178">2023</xref>), and SpikeLM (Xing et al., <xref ref-type="bibr" rid="B144">2024</xref>). Applications regarding equipment safety monitoring mainly include battery health monitoring (Wang et al., <xref ref-type="bibr" rid="B133">2023a</xref>,<xref ref-type="bibr" rid="B125">b</xref>), autonomous vehicle sensors fault diagnosis (Wang and Li, <xref ref-type="bibr" rid="B124">2023</xref>), and bearing fault diagnosis (Xu et al., <xref ref-type="bibr" rid="B146">2022</xref>). Applications in semantic communication mainly tried to mitigate the limitation of transmission bandwidth (Wang M. et al., <xref ref-type="bibr" rid="B126">2023</xref>). Applications in multi-modal information processing currently show up in audio-visual zero-shot learning tasks (Li et al., <xref ref-type="bibr" rid="B71">2023a</xref>,<xref ref-type="bibr" rid="B70">b</xref>).</p>
</sec>
</sec>
<sec>
<title>6.3 Discussion on SNN applications</title>
<p>Deep SNNs have achieved great success in many fields in recent years, but there still exist some limits that need to be addressed. Firstly, although many studies demonstrated that deep SNNs achieved comparable accuracy to their ANN counterparts on many tasks, they still lag behind conventional ANN SOTA, especially for large datasets like ImageNet, which asks for more endeavors. Secondly, many studies claimed that the proposed SNNs consumed much less energy compared to ANN counterparts, through calculating the number of addition operations, without considering the cost of other operations like data movement. Therefore, it is meaningful to deploy well-performed SNNs on neuromorphic chips or corresponding simulators to fully exploit the event-driven nature and measure the actual energy cost. Thirdly, as for applications requiring high processing speed and low power consumption, like robotics, it is promising to adopt neuromorphic vision/audio sensors and neuromorphic processing chips due to their event-driven nature, besides network pruning and weight quantization. Meanwhile, to fully exploit the advantages of events, it deserves more efforts to explore how to directly process neuromorphic sensing events using SNNs, without converting events into frames as current studies usually do. Fourthly, as for transformer-based SNNs used in language or video processing, how to choose the input clip for one simulation step, to reconcile the temporal resolution of the input sequence and the simulation step of SNNs, is worth studying. Last but not least, as SNNs have an additional temporal dimension, how to achieve the speed-accuracy trade-off as humans is a problem worth of study. In other words, how to assign a suitable simulation duration or how to decide when to make a choice, are important questions to realize the balance between computation cost and prediction accuracy.</p>
</sec>
</sec>
<sec id="s7">
<title>7 Future trends and conclusions</title>
<p>This article provides an overview of the current developments in various theories and methods of deep SNNs, including relevant fundamentals, various spiking neuron models, advanced models, and architectures, booming software tools and hardware platforms, as well as applications in various fields. However, there are still many limitations and challenges.</p>
<p>(1) Currently, only a few aspects of the intelligent brains have been applied to instruct the construction and training of SNNs, lacking enough biological plausibility. Therefore, to improve SNNs&#x00027; capability, it is necessary to introduce more types of spiking neurons, rich connection structures, multiscale local-global-cooperative learning rules, system homeostasis, etc., into SNNs to more accurately mimic the cognitive and intelligent characteristics emerging in the brains. For example, it deserves more efforts to train SNNs with self-supervised or unsupervised learning (Zhou Z. et al., <xref ref-type="bibr" rid="B173">2024</xref>), as children mainly receive unlabeled data during growth. Besides, the brain is actually a complex network, thus it is worthy of more effort to study graph SNNs, although some attempts already exist (Li et al., <xref ref-type="bibr" rid="B69">2024</xref>; Yin et al., <xref ref-type="bibr" rid="B154">2024</xref>).</p>
<p>(2) Recent neuroscience studies have found that astrocytes can naturally realize Transformer operations (Kozachkov et al., <xref ref-type="bibr" rid="B61">2023</xref>), which provides a new direction for the improvement of SNNs. In addition, astrocytes have the function of regulating neuronal firing activity and synaptic pruning (Lee et al., <xref ref-type="bibr" rid="B66">2021</xref>; Liu et al., <xref ref-type="bibr" rid="B78">2022</xref>), which provides ideas for the performance improvement and lightweight of SNNs in the future.</p>
<p>(3) Information encoding methods and training algorithms for SNNs are mostly based on average firing rates, lacking the ability to represent temporal dynamics adequately. There should be more exploration of time-dependent information encoding strategies and corresponding training algorithms, to further enhance the spatiotemporal dynamic characteristics of SNNs and strengthen their temporal processing capability.</p>
<p>(4) The training of SNNs mainly employs time-dependent methods, like BPTT, which greatly increases the training cost, compared to conventional DNNs. Thus, there is a need to develop brain-like SNNs that can be trained in parallel, and dedicated software and hardware that support their computation, reducing training time and power consumption.</p>
<p>(5) As there are obstacles to conversion and interaction between different neuromorphic platforms, it is needed to establish a common standard to improve interoperability. Further, more brain-inspired principles or technologies should be incorporated into the design of neuromorphic systems, to enhance the computational performance of the chips, in terms of processing speed and energy efficiency.</p>
<p>(6) Large-scale SNNs are mainly applied to classification tasks. Their potential in handling tasks that need to process continuous input streams, such as videos, languages, events from neuromorphic vision sensors, etc., has not been fully explored. Moreover, the introduction of various neuromorphic sensors and neuromorphic chips into autonomous robotics, cooperating with conventional sensors and processing chips, might be an efficient and effective way to achieve embodied intelligence. Further studies are needed to fully leverage the features and advantages of SNNs.</p>
<p>In summary, studies and applications of SNNs are growing rapidly, but there is still great potential to improve the effectiveness and efficiency of SNNs. Efforts should be made in multiple directions, including model architectures, training algorithms, software frameworks, and hardware platforms, to promote the coordinated progress of models, software, and hardware.</p>
</sec>
<sec sec-type="author-contributions" id="s8">
<title>Author contributions</title>
<p>CZ: Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. HZha: Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. LY: Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. YY: Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. ZZ: Writing &#x02013; review &#x00026; editing. LH: Writing &#x02013; review &#x00026; editing. ZM: Funding acquisition, Project administration, Resources, Supervision, Writing &#x02013; review &#x00026; editing. XF: Writing &#x02013; review &#x00026; editing, Resources. HZho: Writing &#x02013; review &#x00026; editing, Resources. YT: Writing &#x02013; review &#x00026; editing, Resources.</p>
</sec>
</body>
<back>
<sec sec-type="funding-information" id="s9">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research, authorship, and/or publication of this article. The study was funded by the National Natural Science Foundation of China under contracts Nos. 62206141, 62236009, 62332002, 62027804, and 61825101, and the major key project of the Peng Cheng Laboratory (PCL2021A13). Computing support was provided by Pengcheng Cloudbrain.</p>
</sec>
<sec sec-type="COI-statement" id="conf1">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x00027;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Abadi</surname> <given-names>M.</given-names></name> <name><surname>Barham</surname> <given-names>P.</given-names></name> <name><surname>Chen</surname> <given-names>J.</given-names></name> <name><surname>Chen</surname> <given-names>Z.</given-names></name> <name><surname>Davis</surname> <given-names>A.</given-names></name> <name><surname>Dean</surname> <given-names>J.</given-names></name> <etal/></person-group>. (<year>2016</year>). <article-title>&#x0201C;Tensorflow: a system for large-scale machine learning,&#x0201D;</article-title> in <source>Symposium on Operating Systems Design and Implementation (OSDI)</source> (<publisher-loc>Savannah, GA</publisher-loc>), <fpage>265</fpage>&#x02013;<lpage>283</lpage>.</citation>
</ref>
<ref id="B2">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Adrian</surname> <given-names>E. D.</given-names></name> <name><surname>Zotterman</surname> <given-names>Y.</given-names></name></person-group> (<year>1926</year>). <article-title>The impulses produced by sensory nerve endings: Part 3. Impulses set up by touch and pressure</article-title>. <source>J. Physiol</source>. <volume>61</volume>:<fpage>465</fpage>. <pub-id pub-id-type="doi">10.1113/jphysiol.1926.sp002273</pub-id><pub-id pub-id-type="pmid">16993807</pub-id></citation>
</ref>
<ref id="B3">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Akopyan</surname> <given-names>F.</given-names></name> <name><surname>Sawada</surname> <given-names>J.</given-names></name> <name><surname>Cassidy</surname> <given-names>A.</given-names></name> <name><surname>Alvarez-Icaza</surname> <given-names>R.</given-names></name> <name><surname>Arthur</surname> <given-names>J.</given-names></name> <name><surname>Merolla</surname> <given-names>P.</given-names></name> <etal/></person-group>. (<year>2015</year>). <article-title>Truenorth: design and tool flow of a 65 mW 1 million neuron programmable neurosynaptic chip</article-title>. <source>IEEE Transact. Comp. Aided Des. Integr. Circ. Syst</source>. <volume>34</volume>, <fpage>1537</fpage>&#x02013;<lpage>1557</lpage>. <pub-id pub-id-type="doi">10.1109/TCAD.2015.2474396</pub-id></citation>
</ref>
<ref id="B4">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Asgari</surname> <given-names>H.</given-names></name> <name><surname>Maybodi</surname> <given-names>B. M.-N.</given-names></name> <name><surname>Kreiser</surname> <given-names>R.</given-names></name> <name><surname>Sandamirskaya</surname> <given-names>Y.</given-names></name></person-group> (<year>2020</year>). <article-title>Digital multiplier-less spiking neural network architecture of reinforcement learning in a context-dependent task</article-title>. <source>IEEE J. Emerg. Select. Top. Circ. Syst</source>. <volume>10</volume>, <fpage>498</fpage>&#x02013;<lpage>511</lpage>. <pub-id pub-id-type="doi">10.1109/JETCAS.2020.3031040</pub-id></citation>
</ref>
<ref id="B5">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bal</surname> <given-names>M.</given-names></name> <name><surname>Sengupta</surname> <given-names>A.</given-names></name></person-group> (<year>2024</year>). <article-title>pikingbert: distilling bert to train spiking language models using implicit differentiation</article-title>. <source>Proc. AAAI Conf. Artif. Intell</source>. <volume>38</volume>, <fpage>10998</fpage>&#x02013;<lpage>11006</lpage>. <pub-id pub-id-type="doi">10.1609/aaai.v38i10.28975</pub-id></citation>
</ref>
<ref id="B6">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Barchid</surname> <given-names>S.</given-names></name> <name><surname>Allaert</surname> <given-names>B.</given-names></name> <name><surname>Aissaoui</surname> <given-names>A.</given-names></name> <name><surname>Mennesson</surname> <given-names>J.</given-names></name> <name><surname>Djeraba</surname> <given-names>C. C.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;Spiking-fer: spiking neural network for facial expression recognition with event cameras,&#x0201D;</article-title> in <source>International Conference on Content-based Multimedia Indexing (CBMI)</source> (<publisher-loc>Orleans</publisher-loc>), <fpage>1</fpage>&#x02013;<lpage>7</lpage>.</citation>
</ref>
<ref id="B7">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Bellec</surname> <given-names>G.</given-names></name> <name><surname>Salaj</surname> <given-names>D.</given-names></name> <name><surname>Subramoney</surname> <given-names>A.</given-names></name> <name><surname>Legenstein</surname> <given-names>R. A.</given-names></name> <name><surname>Maass</surname> <given-names>W.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Long short-term memory and learning-to-learn in networks of spiking neurons,&#x0201D;</article-title> in <source>NIPS&#x00027;18: Proceedings of the 32nd International Conference on Neural Information Processing Systems, Vol. 31</source> (<publisher-loc>Montreal, QC</publisher-loc>).</citation>
</ref>
<ref id="B8">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bellec</surname> <given-names>G.</given-names></name> <name><surname>Scherr</surname> <given-names>F.</given-names></name> <name><surname>Subramoney</surname> <given-names>A.</given-names></name> <name><surname>Hajek</surname> <given-names>E.</given-names></name> <name><surname>Salaj</surname> <given-names>D.</given-names></name> <name><surname>Legenstein</surname> <given-names>R.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>A solution to the learning dilemma for recurrent networks of spiking neurons</article-title>. <source>Nat. Commun</source>. <volume>11</volume>:<fpage>3625</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-020-17236-y</pub-id><pub-id pub-id-type="pmid">32681001</pub-id></citation>
</ref>
<ref id="B9">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Benjamin</surname> <given-names>B. V.</given-names></name> <name><surname>Gao</surname> <given-names>P.</given-names></name> <name><surname>McQuinn</surname> <given-names>E.</given-names></name> <name><surname>Choudhary</surname> <given-names>S.</given-names></name> <name><surname>Chandrasekaran</surname> <given-names>A. R.</given-names></name> <name><surname>Bussat</surname> <given-names>J.-M.</given-names></name> <etal/></person-group>. (<year>2014</year>). <article-title>Neurogrid: a mixed-analog-digital multichip system for large-scale neural simulations</article-title>. <source>Proc. IEEE</source> <volume>102</volume>, <fpage>699</fpage>&#x02013;<lpage>716</lpage>. <pub-id pub-id-type="doi">10.1109/JPROC.2014.2313565</pub-id></citation>
</ref>
<ref id="B10">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bi</surname> <given-names>G.-Q.</given-names></name> <name><surname>Poo</surname> <given-names>M.-M.</given-names></name></person-group> (<year>1998</year>). <article-title>Synaptic modifications in cultured hippocampal neurons: dependence on spike timing, synaptic strength, and postsynaptic cell type</article-title>. <source>J. Neurosci</source>. <volume>18</volume>, <fpage>10464</fpage>&#x02013;<lpage>10472</lpage>. <pub-id pub-id-type="doi">10.1523/JNEUROSCI.18-24-10464.1998</pub-id><pub-id pub-id-type="pmid">9852584</pub-id></citation>
</ref>
<ref id="B11">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bird</surname> <given-names>G. M.</given-names></name> <name><surname>Polivoda</surname> <given-names>M. E.</given-names></name></person-group> (<year>2021</year>). <article-title>Backpropagation through time for networks with long-term dependencies</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2103.15589</pub-id></citation>
</ref>
<ref id="B12">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Bohnstingl</surname> <given-names>T.</given-names></name> <name><surname>&#x00160;urina</surname> <given-names>A.</given-names></name> <name><surname>Fabre</surname> <given-names>M.</given-names></name> <name><surname>Demira&#x0011F;</surname> <given-names>Y.</given-names></name> <name><surname>Frenkel</surname> <given-names>C.</given-names></name> <name><surname>Payvand</surname> <given-names>M.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>&#x0201C;Biologically-inspired training of spiking recurrent neural networks with neuromorphic hardware,&#x0201D;</article-title> in <source>2022 IEEE 4th International Conference on Artificial Intelligence Circuits and Systems (AICAS)</source> (<publisher-loc>Incheon</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>218</fpage>&#x02013;<lpage>221</lpage>.</citation>
</ref>
<ref id="B13">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bu</surname> <given-names>T.</given-names></name> <name><surname>Fang</surname> <given-names>W.</given-names></name> <name><surname>Ding</surname> <given-names>J.</given-names></name> <name><surname>Dai</surname> <given-names>P.</given-names></name> <name><surname>Yu</surname> <given-names>Z.</given-names></name> <name><surname>Huang</surname> <given-names>T.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Optimal ann-snn conversion for high-accuracy and ultra-low-latency spiking neural networks,&#x0201D;</article-title> in <source>International Conference on Learning Representations (ICLR)</source>.</citation>
</ref>
<ref id="B14">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Bulzomi</surname> <given-names>H.</given-names></name> <name><surname>Schweiker</surname> <given-names>M.</given-names></name> <name><surname>Gruel</surname> <given-names>A.</given-names></name> <name><surname>Martinet</surname> <given-names>J.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;End-to-end neuromorphic lip-reading,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)</source> (<publisher-loc>Vancouver, BC</publisher-loc>), <fpage>4100</fpage>&#x02013;<lpage>4107</lpage>.</citation>
</ref>
<ref id="B15">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cao</surname> <given-names>Y.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Khosla</surname> <given-names>D.</given-names></name></person-group> (<year>2015</year>). <article-title>Spiking deep convolutional neural networks for energy-efficient object recognition</article-title>. <source>Int. J. Comput. Vis</source>. <volume>113</volume>, <fpage>54</fpage>&#x02013;<lpage>66</lpage>. <pub-id pub-id-type="doi">10.1007/s11263-014-0788-3</pub-id></citation>
</ref>
<ref id="B16">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Carion</surname> <given-names>N.</given-names></name> <name><surname>Massa</surname> <given-names>F.</given-names></name> <name><surname>Synnaeve</surname> <given-names>G.</given-names></name> <name><surname>Usunier</surname> <given-names>N.</given-names></name> <name><surname>Kirillov</surname> <given-names>A.</given-names></name> <name><surname>Zagoruyko</surname> <given-names>S.</given-names></name></person-group> (<year>2020</year>). <article-title>End-to-end object detection with transformers</article-title>. <source>Proc. Eur. Conf. Comp. Vis</source>. <volume>12346</volume>, <fpage>213</fpage>&#x02013;<lpage>229</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-030-58452-8_13</pub-id></citation>
</ref>
<ref id="B17">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Castagnetti</surname> <given-names>A.</given-names></name> <name><surname>Pegatoquet</surname> <given-names>A.</given-names></name> <name><surname>Miramond</surname> <given-names>B.</given-names></name></person-group> (<year>2023</year>). <article-title>Spiden: deep spiking neural networks for efficient image denoising</article-title>. <source>Front. Neurosci</source>. <volume>17</volume>:<fpage>1224457</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2023.1224457</pub-id><pub-id pub-id-type="pmid">37638316</pub-id></citation>
</ref>
<ref id="B18">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chakraborty</surname> <given-names>B.</given-names></name> <name><surname>Mukhopadhyay</surname> <given-names>S.</given-names></name></person-group> (<year>2023</year>). <article-title>Heterogeneous recurrent spiking neural network for spatio-temporal classification</article-title>. <source>Front. Neurosci</source>. <volume>17</volume>:<fpage>994517</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2023.994517</pub-id><pub-id pub-id-type="pmid">36793542</pub-id></citation>
</ref>
<ref id="B19">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>D.</given-names></name> <name><surname>Peng</surname> <given-names>P.</given-names></name> <name><surname>Huang</surname> <given-names>T.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2022</year>). <article-title>Deep reinforcement learning with spiking q-learning</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2201.09754</pub-id></citation>
</ref>
<ref id="B20">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>G.</given-names></name> <name><surname>Peng</surname> <given-names>P.</given-names></name> <name><surname>Li</surname> <given-names>G.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2023</year>). <article-title>Training full spike neural networks via auxiliary accumulation pathway</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2301.11929</pub-id></citation>
</ref>
<ref id="B21">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>G. K.</given-names></name> <name><surname>Kumar</surname> <given-names>R.</given-names></name> <name><surname>Sumbul</surname> <given-names>H. E.</given-names></name> <name><surname>Knag</surname> <given-names>P. C.</given-names></name> <name><surname>Krishnamurthy</surname> <given-names>R. K.</given-names></name></person-group> (<year>2018</year>). <article-title>A 4096-neuron 1M-synapse 3.8-pJ/SOP spiking neural network with on-chip STDP learning and sparse weights in 10-nm FinFET CMOS</article-title>. <source>IEEE J. Solid State Circ</source>. <volume>54</volume>, <fpage>992</fpage>&#x02013;<lpage>1002</lpage>. <pub-id pub-id-type="doi">10.1109/JSSC.2018.2884901</pub-id></citation>
</ref>
<ref id="B22">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Com&#x0015F;a</surname> <given-names>I.-M.</given-names></name> <name><surname>Versari</surname> <given-names>L.</given-names></name> <name><surname>Fischbacher</surname> <given-names>T.</given-names></name> <name><surname>Alakuijala</surname> <given-names>J.</given-names></name></person-group> (<year>2021</year>). <article-title>Spiking autoencoders with temporal coding</article-title>. <source>Front. Neurosci</source>. <volume>15</volume>:<fpage>712667</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2021.712667</pub-id><pub-id pub-id-type="pmid">34483829</pub-id></citation>
</ref>
<ref id="B23">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Cordone</surname> <given-names>L.</given-names></name> <name><surname>Miramond</surname> <given-names>B.</given-names></name> <name><surname>Thierion</surname> <given-names>P.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Object detection with spiking neural networks on automotive event data,&#x0201D;</article-title> in <source>International Joint Conference on Neural Networks (IJCNN)</source> (<publisher-loc>Padua</publisher-loc>), <fpage>1</fpage>&#x02013;<lpage>8</lpage>.</citation>
</ref>
<ref id="B24">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cuadrado</surname> <given-names>J.</given-names></name> <name><surname>Ran&#x000E7;on</surname> <given-names>U.</given-names></name> <name><surname>Cottereau</surname> <given-names>B. R.</given-names></name> <name><surname>Barranco</surname> <given-names>F.</given-names></name> <name><surname>Masquelier</surname> <given-names>T.</given-names></name></person-group> (<year>2023</year>). <article-title>Optical flow estimation from event-based cameras and spiking neural networks</article-title>. <source>Front. Neurosci</source>. <volume>17</volume>:<fpage>1160034</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2023.1160034</pub-id><pub-id pub-id-type="pmid">37250425</pub-id></citation>
</ref>
<ref id="B25">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Davies</surname> <given-names>M.</given-names></name></person-group> (<year>2021</year>). <article-title>Taking neuromorphic computing to the next level with Loihi2</article-title>. <source>Intel Labs Loihi</source> <volume>2</volume>, <fpage>1</fpage>&#x02013;<lpage>7</lpage>.</citation>
</ref>
<ref id="B26">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Davies</surname> <given-names>M.</given-names></name> <name><surname>Srinivasa</surname> <given-names>N.</given-names></name> <name><surname>Lin</surname> <given-names>T.-H.</given-names></name> <name><surname>Chinya</surname> <given-names>G.</given-names></name> <name><surname>Cao</surname> <given-names>Y.</given-names></name> <name><surname>Choday</surname> <given-names>S. H.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Loihi: A neuromorphic manycore processor with on-chip learning</article-title>. <source>IEEE Micro</source> <volume>38</volume>, <fpage>82</fpage>&#x02013;<lpage>99</lpage>. <pub-id pub-id-type="doi">10.1109/MM.2018.112130359</pub-id></citation>
</ref>
<ref id="B27">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Deng</surname> <given-names>S.</given-names></name> <name><surname>Li</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>S.</given-names></name> <name><surname>Gu</surname> <given-names>S.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Temporal efficient training of spiking neural network via gradient re-weighting,&#x0201D;</article-title> in <source>International Conference on Learning Representations (ICLR)</source>.</citation>
</ref>
<ref id="B28">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dosovitskiy</surname> <given-names>A.</given-names></name> <name><surname>Beyer</surname> <given-names>L.</given-names></name> <name><surname>Kolesnikov</surname> <given-names>A.</given-names></name> <name><surname>Weissenborn</surname> <given-names>D.</given-names></name> <name><surname>Zhai</surname> <given-names>X.</given-names></name> <name><surname>Unterthiner</surname> <given-names>T.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>&#x0201C;An image is worth 16x16 words: Transformers for image recognition at scale,&#x0201D;</article-title> in <source>International Conference on Learning Representations (ICLR)</source>.</citation>
</ref>
<ref id="B29">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Duan</surname> <given-names>C.</given-names></name> <name><surname>Ding</surname> <given-names>J.</given-names></name> <name><surname>Chen</surname> <given-names>S.</given-names></name> <name><surname>Yu</surname> <given-names>Z.</given-names></name> <name><surname>Huang</surname> <given-names>T.</given-names></name></person-group> (<year>2022</year>). <article-title>Temporal effective batch normalization in spiking neural networks</article-title>. <source>Adv. Neural Inf. Process. Syst</source>. <volume>35</volume>, <fpage>34377</fpage>&#x02013;<lpage>34390</lpage>.</citation>
</ref>
<ref id="B30">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Duwek</surname> <given-names>H. C.</given-names></name> <name><surname>Shalumov</surname> <given-names>A.</given-names></name> <name><surname>Tsur</surname> <given-names>E. E.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;Image reconstruction from neuromorphic event cameras using laplacian-prediction and poisson integration with spiking and artificial neural networks,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition Workshops (CVPR Workshops)</source>, <fpage>1333</fpage>&#x02013;<lpage>1341</lpage>.</citation>
</ref>
<ref id="B31">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Eshraghian</surname> <given-names>J. K.</given-names></name> <name><surname>Ward</surname> <given-names>M.</given-names></name> <name><surname>Neftci</surname> <given-names>E. O.</given-names></name> <name><surname>Wang</surname> <given-names>X.</given-names></name> <name><surname>Lenz</surname> <given-names>G.</given-names></name> <name><surname>Dwivedi</surname> <given-names>G.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Training spiking neural networks using lessons from deep learning</article-title>. <source>Proc. IEEE</source> <volume>111</volume>, <fpage>1016</fpage>&#x02013;<lpage>1054</lpage>. <pub-id pub-id-type="doi">10.1109/JPROC.2023.3308088</pub-id></citation>
</ref>
<ref id="B32">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fang</surname> <given-names>W.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Ding</surname> <given-names>J.</given-names></name> <name><surname>Yu</surname> <given-names>Z.</given-names></name> <name><surname>Masquelier</surname> <given-names>T.</given-names></name> <name><surname>Chen</surname> <given-names>D.</given-names></name> <etal/></person-group>. (<year>2023a</year>). <article-title>Spikingjelly: an open-source machine learning infrastructure platform for spike-based intelligence</article-title>. <source>Sci. Adv</source>. <volume>9</volume>:<fpage>eadi1480</fpage>. <pub-id pub-id-type="doi">10.1126/sciadv.adi1480</pub-id><pub-id pub-id-type="pmid">37801497</pub-id></citation>
</ref>
<ref id="B33">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Fang</surname> <given-names>W.</given-names></name> <name><surname>Yu</surname> <given-names>Z.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Masquelier</surname> <given-names>T.</given-names></name> <name><surname>Huang</surname> <given-names>T.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2021b</year>). <article-title>&#x0201C;Incorporating learnable membrane time constant to enhance learning of spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)</source> (<publisher-loc>Montreal, QC</publisher-loc>), <fpage>2661</fpage>&#x02013;<lpage>2671</lpage>.</citation>
</ref>
<ref id="B34">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Fang</surname> <given-names>W.</given-names></name> <name><surname>Yu</surname> <given-names>Z.</given-names></name> <name><surname>Zhou</surname> <given-names>Z.</given-names></name> <name><surname>Chen</surname> <given-names>D.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <etal/></person-group>. (<year>2023b</year>). <article-title>&#x0201C;Parallel spiking neurons with high efficiency and ability to learn long-term dependencies,&#x0201D;</article-title> in <source>NIPS &#x00027;23: Proceedings of the 37th International Conference on Neural Information Processing Systems, Vol. 36</source> (<publisher-loc>New Orleans, LA</publisher-loc>).</citation>
</ref>
<ref id="B35">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fang</surname> <given-names>W.</given-names></name> <name><surname>Yu</surname> <given-names>Z.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Huang</surname> <given-names>T.</given-names></name> <name><surname>Masquelier</surname> <given-names>T.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2021a</year>). <article-title>Deep residual learning in spiking neural networks</article-title>. <source>Adv. Neural Inf. Process. Syst</source>. <volume>34</volume>, <fpage>21056</fpage>&#x02013;<lpage>21069</lpage>.</citation>
</ref>
<ref id="B36">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Feng</surname> <given-names>L.</given-names></name> <name><surname>Liu</surname> <given-names>Q.</given-names></name> <name><surname>Tang</surname> <given-names>H.</given-names></name> <name><surname>Ma</surname> <given-names>D.</given-names></name> <name><surname>Pan</surname> <given-names>G.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Multi-level firing with spiking ds-resnet: Enabling better and deeper directly-trained spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the Thirty-First International Joint Conference on Artificial Intelligence (IJCAI)</source>, <fpage>2471</fpage>&#x02013;<lpage>2477</lpage>.</citation>
</ref>
<ref id="B37">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Feng</surname> <given-names>Y.</given-names></name> <name><surname>Geng</surname> <given-names>S.</given-names></name> <name><surname>Chu</surname> <given-names>J.</given-names></name> <name><surname>Fu</surname> <given-names>Z.</given-names></name> <name><surname>Hong</surname> <given-names>S.</given-names></name></person-group> (<year>2022</year>). <article-title>Building and training a deep spiking neural network for ecg classification</article-title>. <source>Biomed. Signal Process. Control</source> <volume>77</volume>:<fpage>103749</fpage>. <pub-id pub-id-type="doi">10.1016/j.bspc.2022.103749</pub-id></citation>
</ref>
<ref id="B38">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Frenkel</surname> <given-names>C.</given-names></name> <name><surname>Lefebvre</surname> <given-names>M.</given-names></name> <name><surname>Legat</surname> <given-names>J.-D.</given-names></name> <name><surname>Bol</surname> <given-names>D.</given-names></name></person-group> (<year>2018</year>). <article-title>A 0.086-<italic>mm</italic><sup>2</sup> 12.7-pJ/SOP 64k-synapse 256-neuron online-learning digital spiking neuromorphic processor in 28-nm cmos</article-title>. <source>IEEE Trans. Biomed. Circuits Syst</source>. <volume>13</volume>, <fpage>145</fpage>&#x02013;<lpage>158</lpage>.<pub-id pub-id-type="pmid">30418919</pub-id></citation>
</ref>
<ref id="B39">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Frenkel</surname> <given-names>C.</given-names></name> <name><surname>Legat</surname> <given-names>J.-D.</given-names></name> <name><surname>Bol</surname> <given-names>D.</given-names></name></person-group> (<year>2019</year>). <article-title>Morphic: a 65-nm 738k-synapse/<italic>mm</italic><sup>2</sup> quad-core binary-weight digital neuromorphic processor with stochastic spike-driven online learning</article-title>. <source>IEEE Trans. Biomed. Circuits Syst</source>. <volume>13</volume>, <fpage>999</fpage>&#x02013;<lpage>1010</lpage>. <pub-id pub-id-type="doi">10.1109/TBCAS.2019.2928793</pub-id><pub-id pub-id-type="pmid">31329562</pub-id></citation>
</ref>
<ref id="B40">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Goodman</surname> <given-names>D. F.</given-names></name> <name><surname>Brette</surname> <given-names>R.</given-names></name></person-group> (<year>2009</year>). <article-title>The brian simulator</article-title>. <source>Front. Neurosci</source>. <volume>3</volume>:<fpage>643</fpage>. <pub-id pub-id-type="doi">10.3389/neuro.01.026.2009</pub-id></citation>
</ref>
<ref id="B41">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>N.</given-names></name> <name><surname>Bethge</surname> <given-names>J.</given-names></name> <name><surname>Yang</surname> <given-names>H.</given-names></name> <name><surname>Zhong</surname> <given-names>K.</given-names></name> <name><surname>Ning</surname> <given-names>X.</given-names></name> <name><surname>Meinel</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Boolnet: minimizing the energy consumption of binary neural networks</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2106.06991</pub-id></citation>
</ref>
<ref id="B42">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>W.</given-names></name> <name><surname>Fouda</surname> <given-names>M. E.</given-names></name> <name><surname>Eltawil</surname> <given-names>A. M.</given-names></name> <name><surname>Salama</surname> <given-names>K. N.</given-names></name></person-group> (<year>2021</year>). <article-title>Neural coding in spiking neural networks: a comparative study for robust neuromorphic systems</article-title>. <source>Front. Neurosci</source>. <volume>15</volume>:<fpage>638474</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2021.638474</pub-id><pub-id pub-id-type="pmid">33746705</pub-id></citation>
</ref>
<ref id="B43">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>L.</given-names></name> <name><surname>Liu</surname> <given-names>X.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Huang</surname> <given-names>X.</given-names></name> <etal/></person-group>. (<year>2022a</year>). <article-title>IM-loss: information maximization loss for spiking neural networks</article-title>. <source>Adv. Neural Inf. Process. Syst</source>. <volume>35</volume>, <fpage>156</fpage>&#x02013;<lpage>166</lpage>. <pub-id pub-id-type="doi">10.20944/preprints202312.1318.v1</pub-id></citation>
</ref>
<ref id="B44">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Tong</surname> <given-names>X.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>L.</given-names></name> <name><surname>Liu</surname> <given-names>X.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <etal/></person-group>. (<year>2022b</year>). <article-title>&#x0201C;RecDis-SNN: rectifying membrane potential distribution for directly training spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)</source> (<publisher-loc>New Orleans, LA</publisher-loc>), <fpage>326</fpage>&#x02013;<lpage>335</lpage>.</citation>
</ref>
<ref id="B45">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Peng</surname> <given-names>W.</given-names></name> <name><surname>Liu</surname> <given-names>X.</given-names></name> <name><surname>Zhang</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2023b</year>). <article-title>&#x0201C;Membrane potential batch normalization for spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)</source> (<publisher-loc>Paris</publisher-loc>), <fpage>19420</fpage>&#x02013;<lpage>19430</lpage>.</citation>
</ref>
<ref id="B46">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Liu</surname> <given-names>X.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>L.</given-names></name> <name><surname>Peng</surname> <given-names>W.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2023a</year>). <article-title>&#x0201C;RMP-loss: Regularizing membrane potential distribution for spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)</source> (<publisher-loc>Paris</publisher-loc>), <fpage>17391</fpage>&#x02013;<lpage>17401</lpage>.</citation>
</ref>
<ref id="B47">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hamilton</surname> <given-names>K. E.</given-names></name> <name><surname>Mintz</surname> <given-names>T. M.</given-names></name> <name><surname>Schuman</surname> <given-names>C. D.</given-names></name></person-group> (<year>2019</year>). <article-title>Spike-based primitives for graph algorithms</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1903.10574</pub-id></citation>
</ref>
<ref id="B48">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hazan</surname> <given-names>H.</given-names></name> <name><surname>Saunders</surname> <given-names>D. J.</given-names></name> <name><surname>Khan</surname> <given-names>H.</given-names></name> <name><surname>Patel</surname> <given-names>D.</given-names></name> <name><surname>Sanghavi</surname> <given-names>D. T.</given-names></name> <name><surname>Siegelmann</surname> <given-names>H. T.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>BINDSNET: a machine learning-oriented spiking neural networks library in python</article-title>. <source>Front. Neuroinform</source>. <volume>12</volume>:<fpage>89</fpage>. <pub-id pub-id-type="doi">10.3389/fninf.2018.00089</pub-id><pub-id pub-id-type="pmid">30631269</pub-id></citation>
</ref>
<ref id="B49">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hines</surname> <given-names>M. L.</given-names></name> <name><surname>Carnevale</surname> <given-names>N. T.</given-names></name></person-group> (<year>1997</year>). <article-title>The neuron simulation environment</article-title>. <source>Neural Comput</source>. <volume>9</volume>, <fpage>1179</fpage>&#x02013;<lpage>1209</lpage>. <pub-id pub-id-type="doi">10.1162/neco.1997.9.6.1179</pub-id><pub-id pub-id-type="pmid">9248061</pub-id></citation>
</ref>
<ref id="B50">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hu</surname> <given-names>Y.</given-names></name> <name><surname>Tang</surname> <given-names>H.</given-names></name> <name><surname>Pan</surname> <given-names>G.</given-names></name></person-group> (<year>2021a</year>). <article-title>Spiking deep residual networks</article-title>. <source>IEEE Transact. Neural Netw. Learn. Syst</source>. <volume>34</volume>, <fpage>5200</fpage>&#x02013;<lpage>5205</lpage>. <pub-id pub-id-type="doi">10.1109/TNNLS.2021.3119238</pub-id></citation>
</ref>
<ref id="B51">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hu</surname> <given-names>Y.</given-names></name> <name><surname>Wu</surname> <given-names>Y.</given-names></name> <name><surname>Deng</surname> <given-names>L.</given-names></name> <name><surname>Li</surname> <given-names>G.</given-names></name></person-group> (<year>2021b</year>). <article-title>Advancing residual learning towards powerful deep spiking neural networks</article-title>. <source>arXiv</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2112.08954</pub-id></citation>
</ref>
<ref id="B52">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Huang</surname> <given-names>L.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Zhou</surname> <given-names>H.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2023b</year>). <article-title>Deep recurrent spiking neural networks capture both static and dynamic representations of the visual cortex under movie stimuli</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2306.01354</pub-id></citation>
</ref>
<ref id="B53">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Huang</surname> <given-names>L.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Yu</surname> <given-names>L.</given-names></name> <name><surname>Zhou</surname> <given-names>H.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2023a</year>). <article-title>&#x0201C;Deep spiking neural networks with high representation similarity model visual pathways of macaque and mouse,&#x0201D;</article-title> in <source>Proceedings of the AAAI Conference on Artificial Intelligence (AAAI)</source> (<publisher-loc>Washington, DC</publisher-loc>), <fpage>31</fpage>&#x02013;<lpage>39</lpage>.</citation>
</ref>
<ref id="B54">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hunsberger</surname> <given-names>E.</given-names></name> <name><surname>Eliasmith</surname> <given-names>C.</given-names></name></person-group> (<year>2015</year>). <article-title>Spiking deep networks with lif neurons</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1510.08829</pub-id></citation>
</ref>
<ref id="B55">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jiang</surname> <given-names>C.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name></person-group> (<year>2023</year>). <article-title>KLIF: an optimized spiking neuron unit for tuning surrogate gradient slope and membrane potential</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2302.09238</pub-id></citation>
</ref>
<ref id="B56">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kamata</surname> <given-names>H.</given-names></name> <name><surname>Mukuta</surname> <given-names>Y.</given-names></name> <name><surname>Harada</surname> <given-names>T.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Fully spiking variational autoencoder,&#x0201D;</article-title> in <source>Proceedings of the AAAI Conference on Artificial Intelligence (AAAI)</source>, <fpage>7059</fpage>&#x02013;<lpage>7067</lpage>.</citation>
</ref>
<ref id="B57">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Kim</surname> <given-names>S.</given-names></name> <name><surname>Park</surname> <given-names>S.</given-names></name> <name><surname>Na</surname> <given-names>B.</given-names></name> <name><surname>Yoon</surname> <given-names>S.</given-names></name></person-group> (<year>2020</year>). <article-title>&#x0201C;Spiking-yolo: spiking neural network for energy-efficient object detection,&#x0201D;</article-title> in <source>Proceedings of the AAAI Conference on Artificial Intelligence (AAAI)</source> (<publisher-loc>New York, NY</publisher-loc>), <fpage>11270</fpage>&#x02013;<lpage>11277</lpage>.</citation>
</ref>
<ref id="B58">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kim</surname> <given-names>Y.</given-names></name> <name><surname>Chough</surname> <given-names>J.</given-names></name> <name><surname>Panda</surname> <given-names>P.</given-names></name></person-group> (<year>2022</year>). <article-title>Beyond classification: directly training spiking neural networks for semantic segmentation</article-title>. <source>Neuromorp. Comp. Eng</source>. <volume>2</volume>:<fpage>044015</fpage>. <pub-id pub-id-type="doi">10.1088/2634-4386/ac9b86</pub-id></citation>
</ref>
<ref id="B59">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kim</surname> <given-names>Y.</given-names></name> <name><surname>Panda</surname> <given-names>P.</given-names></name></person-group> (<year>2021</year>). <article-title>Revisiting batch normalization for training low-latency deep spiking neural networks from scratch</article-title>. <source>Front. Neurosci</source>. <volume>15</volume>:<fpage>773954</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2021.773954</pub-id><pub-id pub-id-type="pmid">34955725</pub-id></citation>
</ref>
<ref id="B60">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Kosta</surname> <given-names>A. K.</given-names></name> <name><surname>Roy</surname> <given-names>K.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;Adaptive-spikenet: event-based optical flow estimation using spiking neural networks with learnable neuronal dynamics,&#x0201D;</article-title> in <source>International Conference on Robotics and Automation (ICRA)</source> (<publisher-loc>London</publisher-loc>), <fpage>6021</fpage>&#x02013;<lpage>6027</lpage>.</citation>
</ref>
<ref id="B61">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kozachkov</surname> <given-names>L.</given-names></name> <name><surname>Kastanenka</surname> <given-names>K. V.</given-names></name> <name><surname>Krotov</surname> <given-names>D.</given-names></name></person-group> (<year>2023</year>). <article-title>Building transformers from neurons and astrocytes</article-title>. <source>Proc. Nat. Acad. Sci. U. S. A</source>. <volume>120</volume>:<fpage>e2219150120</fpage>. <pub-id pub-id-type="doi">10.1073/pnas.2219150120</pub-id></citation>
</ref>
<ref id="B62">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kuang</surname> <given-names>Y.</given-names></name> <name><surname>Cui</surname> <given-names>X.</given-names></name> <name><surname>Zhong</surname> <given-names>Y.</given-names></name> <name><surname>Liu</surname> <given-names>K.</given-names></name> <name><surname>Zou</surname> <given-names>C.</given-names></name> <name><surname>Dai</surname> <given-names>Z.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>A 64K-neuron 64M-1b-synapse 2.64 pJ/SOP neuromorphic chip with all memory on chip for spike-based models in 65nm CMOS</article-title>. <source>IEEE Transact. Circ. Syst. II</source> <volume>68</volume>, <fpage>2655</fpage>&#x02013;<lpage>2659</lpage>. <pub-id pub-id-type="doi">10.1109/TCSII.2021.3052172</pub-id></citation>
</ref>
<ref id="B63">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Kugele</surname> <given-names>A.</given-names></name> <name><surname>Pfeil</surname> <given-names>T.</given-names></name> <name><surname>Pfeiffer</surname> <given-names>M.</given-names></name> <name><surname>Chicca</surname> <given-names>E.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;Hybrid SNN-ANN: energy-efficient classification and object detection for event-based vision,&#x0201D;</article-title> in <source>DAGM German Conference on Pattern Recognition (GCPR)</source> (<publisher-loc>Bonn</publisher-loc>), <fpage>297</fpage>&#x02013;<lpage>312</lpage>.</citation>
</ref>
<ref id="B64">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lee</surname> <given-names>C.</given-names></name> <name><surname>Kosta</surname> <given-names>A. K.</given-names></name> <name><surname>Zhu</surname> <given-names>A. Z.</given-names></name> <name><surname>Chaney</surname> <given-names>K.</given-names></name> <name><surname>Daniilidis</surname> <given-names>K.</given-names></name> <name><surname>Roy</surname> <given-names>K.</given-names></name></person-group> (<year>2020a</year>). <article-title>Spike-flownet: event-based optical flow estimation with energy-efficient hybrid neural networks</article-title>. <source>Proc. Eur. Conf. Comp. Vis</source>. <volume>12374</volume>, <fpage>366</fpage>&#x02013;<lpage>382</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-030-58526-6_22</pub-id></citation>
</ref>
<ref id="B65">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lee</surname> <given-names>C.</given-names></name> <name><surname>Sarwar</surname> <given-names>S. S.</given-names></name> <name><surname>Panda</surname> <given-names>P.</given-names></name> <name><surname>Srinivasan</surname> <given-names>G.</given-names></name> <name><surname>Roy</surname> <given-names>K.</given-names></name></person-group> (<year>2020b</year>). <article-title>Enabling spike-based backpropagation for training deep neural network architectures</article-title>. <source>Front. Neurosci</source>. <volume>14</volume>:<fpage>119</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2020.00119</pub-id><pub-id pub-id-type="pmid">32180697</pub-id></citation>
</ref>
<ref id="B66">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lee</surname> <given-names>J.-H.</given-names></name> <name><surname>Kim</surname> <given-names>J.-Y.</given-names></name> <name><surname>Noh</surname> <given-names>S.</given-names></name> <name><surname>Lee</surname> <given-names>H.</given-names></name> <name><surname>Lee</surname> <given-names>S. Y.</given-names></name> <name><surname>Mun</surname> <given-names>J. Y.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Astrocytes phagocytose adult hippocampal synapses for circuit homeostasis</article-title>. <source>Nature</source> <volume>590</volume>, <fpage>612</fpage>&#x02013;<lpage>617</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-020-03060-3</pub-id><pub-id pub-id-type="pmid">33361813</pub-id></citation>
</ref>
<ref id="B67">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lee</surname> <given-names>J. H.</given-names></name> <name><surname>Delbruck</surname> <given-names>T.</given-names></name> <name><surname>Pfeiffer</surname> <given-names>M.</given-names></name></person-group> (<year>2016</year>). <article-title>Training deep spiking neural networks using backpropagation</article-title>. <source>Front. Neurosci</source>. <volume>10</volume>:<fpage>508</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2016.00508</pub-id><pub-id pub-id-type="pmid">27877107</pub-id></citation>
</ref>
<ref id="B68">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2021</year>). <article-title>Recent advances in neuromorphic vision sensors: a survey</article-title>. <source>Chin. J. Comp</source>. <volume>44</volume>, <fpage>1258</fpage>&#x02013;<lpage>1286</lpage>.</citation>
</ref>
<ref id="B69">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Wu</surname> <given-names>R.</given-names></name> <name><surname>Zhu</surname> <given-names>Z.</given-names></name> <name><surname>Wang</surname> <given-names>B.</given-names></name> <name><surname>Meng</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>&#x0201C;A graph is worth 1-bit spikes: When graph contrastive learning meets spiking neural networks,&#x0201D;</article-title> in <source>International Conference on Learning Representations (ICLR)</source>.</citation>
</ref>
<ref id="B70">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>W.</given-names></name> <name><surname>Zhao</surname> <given-names>X.-L.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Wang</surname> <given-names>X.</given-names></name> <name><surname>Fan</surname> <given-names>X.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2023b</year>). <article-title>&#x0201C;Motion-decoupled spiking transformer for audio-visual zero-shot learning,&#x0201D;</article-title> in <source>Proceedings of the International Conference on Multimedia (MM)</source>, <fpage>3994</fpage>&#x02013;<lpage>4002</lpage>.</citation>
</ref>
<ref id="B71">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>W.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Deng</surname> <given-names>L.-J.</given-names></name> <name><surname>Man</surname> <given-names>H.</given-names></name> <name><surname>Fan</surname> <given-names>X.</given-names></name></person-group> (<year>2023a</year>). <article-title>&#x0201C;Modality-fusion spiking transformer network for audio-visual zero-shot learning,&#x0201D;</article-title> in <source>International Conference on Multimedia and Expo (ICME)</source> (<publisher-loc>Brisbane, QLD</publisher-loc>), <fpage>426</fpage>&#x02013;<lpage>431</lpage>.</citation>
</ref>
<ref id="B72">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>X.</given-names></name> <name><surname>Zhang</surname> <given-names>X.</given-names></name> <name><surname>Yi</surname> <given-names>X.</given-names></name> <name><surname>Liu</surname> <given-names>D.</given-names></name> <name><surname>Wang</surname> <given-names>H.</given-names></name> <name><surname>Zhang</surname> <given-names>B.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Review of medical data analysis based on spiking neural networks</article-title>. <source>Proc. Comput. Sci</source>. <volume>221</volume>, <fpage>1527</fpage>&#x02013;<lpage>1538</lpage>. <pub-id pub-id-type="doi">10.1016/j.procs.2023.08.138</pub-id></citation>
</ref>
<ref id="B73">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>Y.</given-names></name> <name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>S.</given-names></name> <name><surname>Deng</surname> <given-names>S.</given-names></name> <name><surname>Hai</surname> <given-names>Y.</given-names></name> <name><surname>Gu</surname> <given-names>S.</given-names></name></person-group> (<year>2021</year>). <article-title>Differentiable spike: rethinking gradient-descent for training spiking neural networks</article-title>. <source>Adv. Neural Inf. Process. Syst</source>. <volume>34</volume>, <fpage>23426</fpage>&#x02013;<lpage>23439</lpage>.</citation>
</ref>
<ref id="B74">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>Y.</given-names></name> <name><surname>He</surname> <given-names>X.</given-names></name> <name><surname>Dong</surname> <given-names>Y.</given-names></name> <name><surname>Kong</surname> <given-names>Q.</given-names></name> <name><surname>Zeng</surname> <given-names>Y.</given-names></name></person-group> (<year>2022</year>). <article-title>Spike calibration: fast and accurate conversion of spiking neural network for object detection and segmentation</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.24963/ijcai.2022/345</pub-id></citation>
</ref>
<ref id="B75">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Lian</surname> <given-names>S.</given-names></name> <name><surname>Shen</surname> <given-names>J.</given-names></name> <name><surname>Liu</surname> <given-names>Q.</given-names></name> <name><surname>Wang</surname> <given-names>Z.</given-names></name> <name><surname>Yan</surname> <given-names>R.</given-names></name> <name><surname>Tang</surname> <given-names>H.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;Learnable surrogate gradient for direct training spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the Thirty-Second International Joint Conference on Artificial Intelligence (IJCAI)</source> (<publisher-loc>Macao</publisher-loc>), <fpage>3002</fpage>&#x02013;<lpage>3010</lpage>.</citation>
</ref>
<ref id="B76">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liang</surname> <given-names>J.</given-names></name> <name><surname>Li</surname> <given-names>R.</given-names></name> <name><surname>Wang</surname> <given-names>C.</given-names></name> <name><surname>Zhang</surname> <given-names>R.</given-names></name> <name><surname>Yue</surname> <given-names>K.</given-names></name> <name><surname>Li</surname> <given-names>W.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>A spiking neural network based on retinal ganglion cells for automatic burn image segmentation</article-title>. <source>Entropy</source> <volume>24</volume>:<fpage>1526</fpage>. <pub-id pub-id-type="doi">10.3390/e24111526</pub-id><pub-id pub-id-type="pmid">36359618</pub-id></citation>
</ref>
<ref id="B77">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liang</surname> <given-names>L.</given-names></name> <name><surname>Qu</surname> <given-names>Z.</given-names></name> <name><surname>Chen</surname> <given-names>Z.</given-names></name> <name><surname>Tu</surname> <given-names>F.</given-names></name> <name><surname>Wu</surname> <given-names>Y.</given-names></name> <name><surname>Deng</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>H2learn: high-efficiency learning accelerator for high-accuracy spiking neural networks</article-title>. <source>IEEE Transac. Comp. Aided Des. Integr. Circ. Syst</source>. <volume>41</volume>, <fpage>4782</fpage>&#x02013;<lpage>4796</lpage>. <pub-id pub-id-type="doi">10.1109/TCAD.2021.3138347</pub-id></citation>
</ref>
<ref id="B78">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>J.-H.</given-names></name> <name><surname>Zhang</surname> <given-names>M.</given-names></name> <name><surname>Wang</surname> <given-names>Q.</given-names></name> <name><surname>Wu</surname> <given-names>D.-Y.</given-names></name> <name><surname>Jie</surname> <given-names>W.</given-names></name> <name><surname>Hu</surname> <given-names>N.-Y.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Distinct roles of astroglia and neurons in synaptic plasticity and memory</article-title>. <source>Mol. Psychiatry</source> <volume>27</volume>, <fpage>873</fpage>&#x02013;<lpage>885</lpage>. <pub-id pub-id-type="doi">10.1038/s41380-021-01332-6</pub-id><pub-id pub-id-type="pmid">34642458</pub-id></citation>
</ref>
<ref id="B79">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>M.</given-names></name> <name><surname>Wen</surname> <given-names>R.</given-names></name> <name><surname>Chen</surname> <given-names>H.</given-names></name></person-group> (<year>2023</year>). <article-title>Spiking-diffusion: vector quantized discrete diffusion model with spiking neural networks</article-title>. <source>arXiv</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2308.10187</pub-id></citation>
</ref>
<ref id="B80">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>Z.</given-names></name> <name><surname>Shen</surname> <given-names>Z.</given-names></name> <name><surname>Savvides</surname> <given-names>M.</given-names></name> <name><surname>Cheng</surname> <given-names>K.-T.</given-names></name></person-group> (<year>2020</year>). <article-title>ReActNet: towards precise binary neural network with generalized activation functions</article-title>. <source>Proc. Eur. Conf. Comp. Vis</source>. <volume>12359</volume>, <fpage>143</fpage>&#x02013;<lpage>159</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-030-58568-6_9</pub-id></citation>
</ref>
<ref id="B81">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>Z.</given-names></name> <name><surname>Wu</surname> <given-names>B.</given-names></name> <name><surname>Luo</surname> <given-names>W.</given-names></name> <name><surname>Yang</surname> <given-names>X.</given-names></name> <name><surname>Liu</surname> <given-names>W.</given-names></name> <name><surname>Cheng</surname> <given-names>K.-T.</given-names></name></person-group> (<year>2018</year>). <article-title>Bi-Real Net: enhancing the performance of 1-bit CNNS with improved representational capability and advanced training algorithm</article-title>. <source>Proc. Eur. Conf. Comp. Vis</source>. <volume>11219</volume>, <fpage>747</fpage>&#x02013;<lpage>763</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-030-01267-0_44</pub-id></citation>
</ref>
<ref id="B82">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>Z.</given-names></name> <name><surname>Lin</surname> <given-names>Y.</given-names></name> <name><surname>Cao</surname> <given-names>Y.</given-names></name> <name><surname>Hu</surname> <given-names>H.</given-names></name> <name><surname>Wei</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>Z.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>&#x0201C;Swin transformer: Hierarchical vision transformer using shifted windows,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)</source>, <fpage>10012</fpage>&#x02013;<lpage>10022</lpage>.</citation>
</ref>
<ref id="B83">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Luo</surname> <given-names>Y.</given-names></name> <name><surname>Shen</surname> <given-names>H.</given-names></name> <name><surname>Cao</surname> <given-names>X.</given-names></name> <name><surname>Wang</surname> <given-names>T.</given-names></name> <name><surname>Feng</surname> <given-names>Q.</given-names></name> <name><surname>Tan</surname> <given-names>Z.</given-names></name></person-group> (<year>2022</year>). <article-title>Conversion of siamese networks to spiking neural networks for energy-efficient object tracking</article-title>. <source>Neur. Comp. Appl</source>. <volume>34</volume>, <fpage>9967</fpage>&#x02013;<lpage>9982</lpage>. <pub-id pub-id-type="doi">10.1007/s00521-022-06984-1</pub-id></citation>
</ref>
<ref id="B84">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lv</surname> <given-names>C.</given-names></name> <name><surname>Li</surname> <given-names>T.</given-names></name> <name><surname>Xu</surname> <given-names>J.</given-names></name> <name><surname>Gu</surname> <given-names>C.</given-names></name> <name><surname>Ling</surname> <given-names>Z.</given-names></name> <name><surname>Zhang</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>SpikeBERT: a language spikformer trained with two-stage knowledge distillation from bert</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2308.15122</pub-id></citation>
</ref>
<ref id="B85">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ma</surname> <given-names>D.</given-names></name> <name><surname>Jin</surname> <given-names>X.</given-names></name> <name><surname>Sun</surname> <given-names>S.</given-names></name> <name><surname>Li</surname> <given-names>Y.</given-names></name> <name><surname>Wu</surname> <given-names>X.</given-names></name> <name><surname>Hu</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>Darwin3: a large-scale neuromorphic chip with a novel isa and on-chip learning</article-title>. <source>Natl. Sci. Rev</source>. <volume>11</volume>:<fpage>nwae102</fpage>. <pub-id pub-id-type="doi">10.1093/nsr/nwae102</pub-id><pub-id pub-id-type="pmid">38689713</pub-id></citation>
</ref>
<ref id="B86">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ma</surname> <given-names>D.</given-names></name> <name><surname>Shen</surname> <given-names>J.</given-names></name> <name><surname>Gu</surname> <given-names>Z.</given-names></name> <name><surname>Zhang</surname> <given-names>M.</given-names></name> <name><surname>Zhu</surname> <given-names>X.</given-names></name> <name><surname>Xu</surname> <given-names>X.</given-names></name> <etal/></person-group>. (<year>2017</year>). <article-title>Darwin: a neuromorphic hardware co-processor based on spiking neural networks</article-title>. <source>J. Syst. Arch</source>. <volume>77</volume>, <fpage>43</fpage>&#x02013;<lpage>51</lpage>. <pub-id pub-id-type="doi">10.1016/j.sysarc.2017.01.003</pub-id></citation>
</ref>
<ref id="B87">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Ma</surname> <given-names>G.</given-names></name> <name><surname>Jiang</surname> <given-names>R.</given-names></name> <name><surname>Yan</surname> <given-names>R.</given-names></name> <name><surname>Tang</surname> <given-names>H.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;Temporal conditioning spiking latent variable models of the neural response to natural visual scenes,&#x0201D;</article-title> in <source>NIPS &#x00027;23: Proceedings of the 37th International Conference on Neural Information Processing Systems, Vol. 36</source> (<publisher-loc>New Orleans, LA</publisher-loc>).</citation>
</ref>
<ref id="B88">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Maass</surname> <given-names>W.</given-names></name></person-group> (<year>1997</year>). <article-title>Networks of spiking neurons: the third generation of neural network models</article-title>. <source>Neural Netw</source>. <volume>10</volume>, <fpage>1659</fpage>&#x02013;<lpage>1671</lpage>. <pub-id pub-id-type="doi">10.1016/S0893-6080(97)00011-7</pub-id></citation>
</ref>
<ref id="B89">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Meng</surname> <given-names>Q.</given-names></name> <name><surname>Xiao</surname> <given-names>M.</given-names></name> <name><surname>Yan</surname> <given-names>S.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Lin</surname> <given-names>Z.</given-names></name> <name><surname>Luo</surname> <given-names>Z.-Q.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Training high-performance low-latency spiking neural networks by differentiation on spike representation,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)</source> (<publisher-loc>New Orleans, LA</publisher-loc>), <fpage>12444</fpage>&#x02013;<lpage>12453</lpage>.</citation>
</ref>
<ref id="B90">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Meng</surname> <given-names>Q.</given-names></name> <name><surname>Xiao</surname> <given-names>M.</given-names></name> <name><surname>Yan</surname> <given-names>S.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Lin</surname> <given-names>Z.</given-names></name> <name><surname>Luo</surname> <given-names>Z.-Q.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;Towards memory- and time-efficient backpropagation for training spiking neural networks,&#x0201D;</article-title> in <source>2023 IEEE/CVF International Conference on Computer Vision (ICCV)</source> (<publisher-loc>Paris</publisher-loc>), <fpage>6143</fpage>&#x02013;<lpage>6153</lpage>.</citation>
</ref>
<ref id="B91">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Modha</surname> <given-names>D. S.</given-names></name> <name><surname>Akopyan</surname> <given-names>F.</given-names></name> <name><surname>Andreopoulos</surname> <given-names>A.</given-names></name> <name><surname>Appuswamy</surname> <given-names>R.</given-names></name> <name><surname>Arthur</surname> <given-names>J. V.</given-names></name> <name><surname>Cassidy</surname> <given-names>A. S.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>&#x0201C;IBM northpole neural inference machine,&#x0201D;</article-title> in <source>2023 IEEE Hot Chips 35 Symposium (HCS)</source> (<publisher-loc>California, MA</publisher-loc>: <publisher-name>IEEE Computer Society</publisher-name>), <fpage>1</fpage>&#x02013;<lpage>58</lpage>.</citation>
</ref>
<ref id="B92">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Moradi</surname> <given-names>S.</given-names></name> <name><surname>Qiao</surname> <given-names>N.</given-names></name> <name><surname>Stefanini</surname> <given-names>F.</given-names></name> <name><surname>Indiveri</surname> <given-names>G.</given-names></name></person-group> (<year>2017</year>). <article-title>A scalable multicore architecture with heterogeneous memory structures for dynamic neuromorphic asynchronous processors (dynaps)</article-title>. <source>IEEE Trans. Biomed. Circuits Syst</source>. <volume>12</volume>, <fpage>106</fpage>&#x02013;<lpage>122</lpage>. <pub-id pub-id-type="doi">10.1109/TBCAS.2017.2759700</pub-id></citation>
</ref>
<ref id="B93">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mozafari</surname> <given-names>M.</given-names></name> <name><surname>Ganjtabesh</surname> <given-names>M.</given-names></name> <name><surname>Nowzari-Dalini</surname> <given-names>A.</given-names></name> <name><surname>Masquelier</surname> <given-names>T.</given-names></name></person-group> (<year>2019</year>). <article-title>SpykeTorch: efficient simulation of convolutional spiking neural networks with at most one spike per neuron</article-title>. <source>Front. Neurosci</source>. <volume>13</volume>:<fpage>625</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2019.00625</pub-id><pub-id pub-id-type="pmid">31354403</pub-id></citation>
</ref>
<ref id="B94">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Neftci</surname> <given-names>E. O.</given-names></name> <name><surname>Mostafa</surname> <given-names>H.</given-names></name> <name><surname>Zenke</surname> <given-names>F.</given-names></name></person-group> (<year>2019</year>). <article-title>Surrogate gradient learning in spiking neural networks: Bringing the power of gradient-based optimization to spiking neural networks</article-title>. <source>IEEE Signal Process. Mag</source>. <volume>36</volume>, <fpage>51</fpage>&#x02013;<lpage>63</lpage>. <pub-id pub-id-type="doi">10.1109/MSP.2019.2931595</pub-id></citation>
</ref>
<ref id="B95">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Painkras</surname> <given-names>E.</given-names></name> <name><surname>Plana</surname> <given-names>L. A.</given-names></name> <name><surname>Garside</surname> <given-names>J.</given-names></name> <name><surname>Temple</surname> <given-names>S.</given-names></name> <name><surname>Galluppi</surname> <given-names>F.</given-names></name> <name><surname>Patterson</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2013</year>). <article-title>SpiNNaker: a 1-W 18-core system-on-chip for massively-parallel neural network simulation</article-title>. <source>IEEE J. Solid State Circ</source>. <volume>48</volume>, <fpage>1943</fpage>&#x02013;<lpage>1953</lpage>. <pub-id pub-id-type="doi">10.1109/JSSC.2013.2259038</pub-id></citation>
</ref>
<ref id="B96">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Panda</surname> <given-names>P.</given-names></name> <name><surname>Srinivasa</surname> <given-names>N.</given-names></name></person-group> (<year>2018</year>). <article-title>Learning to recognize actions from limited training examples using a recurrent spiking neural model</article-title>. <source>Front. Neurosci</source>. <volume>12</volume>:<fpage>126</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2018.00126</pub-id><pub-id pub-id-type="pmid">29551962</pub-id></citation>
</ref>
<ref id="B97">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Parameshwara</surname> <given-names>C. M.</given-names></name> <name><surname>Li</surname> <given-names>S.</given-names></name> <name><surname>Ferm&#x000FC;ller</surname> <given-names>C.</given-names></name> <name><surname>Sanket</surname> <given-names>N. J.</given-names></name> <name><surname>Evanusa</surname> <given-names>M. S.</given-names></name> <name><surname>Aloimonos</surname> <given-names>Y.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;Spikems: deep spiking neural network for motion segmentation,&#x0201D;</article-title> in <source>International Conference on Intelligent Robots and Systems (IROS)</source>, <fpage>3414</fpage>&#x02013;<lpage>3420</lpage>.</citation>
</ref>
<ref id="B98">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Park</surname> <given-names>S.</given-names></name> <name><surname>Kim</surname> <given-names>S.</given-names></name> <name><surname>Na</surname> <given-names>B.</given-names></name> <name><surname>Yoon</surname> <given-names>S.</given-names></name></person-group> (<year>2020</year>). <article-title>&#x0201C;T2FSNN: deep spiking neural networks with time-to-first-spike coding,&#x0201D;</article-title> in <source>ACM/IEEE Design Automation Conference (DAC)</source> (<publisher-loc>San Francisco, CA</publisher-loc>), <fpage>1</fpage>&#x02013;<lpage>6</lpage>.</citation>
</ref>
<ref id="B99">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Paszke</surname> <given-names>A.</given-names></name> <name><surname>Gross</surname> <given-names>S.</given-names></name> <name><surname>Massa</surname> <given-names>F.</given-names></name> <name><surname>Lerer</surname> <given-names>A.</given-names></name> <name><surname>Bradbury</surname> <given-names>J.</given-names></name> <name><surname>Chanan</surname> <given-names>G.</given-names></name> <etal/></person-group>. (<year>2019</year>). <article-title>&#x0201C;PyTorch: an imperative style, high-performance deep learning library,&#x0201D;</article-title> in <source>Advances in Neural Information Processing Systems (NeurIPS), Vol. 32</source> (<publisher-loc>British Columbia</publisher-loc>).</citation>
</ref>
<ref id="B100">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pehle</surname> <given-names>C.</given-names></name> <name><surname>Billaudelle</surname> <given-names>S.</given-names></name> <name><surname>Cramer</surname> <given-names>B.</given-names></name> <name><surname>Kaiser</surname> <given-names>J.</given-names></name> <name><surname>Schreiber</surname> <given-names>K.</given-names></name> <name><surname>Stradmann</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>The brainscales-2 accelerated neuromorphic system with hybrid plasticity</article-title>. <source>Front. Neurosci</source>. <volume>16</volume>:<fpage>795876</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2022.795876</pub-id><pub-id pub-id-type="pmid">35281488</pub-id></citation>
</ref>
<ref id="B101">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pei</surname> <given-names>J.</given-names></name> <name><surname>Deng</surname> <given-names>L.</given-names></name> <name><surname>Song</surname> <given-names>S.</given-names></name> <name><surname>Zhao</surname> <given-names>M.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Wu</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2019</year>). <article-title>Towards artificial general intelligence with hybrid tianjic chip architecture</article-title>. <source>Nature</source> <volume>572</volume>, <fpage>106</fpage>&#x02013;<lpage>111</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-019-1424-8</pub-id><pub-id pub-id-type="pmid">31367028</pub-id></citation>
</ref>
<ref id="B102">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Qasim Gilani</surname> <given-names>S.</given-names></name> <name><surname>Syed</surname> <given-names>T.</given-names></name> <name><surname>Umair</surname> <given-names>M.</given-names></name> <name><surname>Marques</surname> <given-names>O.</given-names></name></person-group> (<year>2023</year>). <article-title>Skin cancer classification using deep spiking neural network</article-title>. <source>J. Digit. Imaging</source> <volume>36</volume>, <fpage>1137</fpage>&#x02013;<lpage>1147</lpage>. <pub-id pub-id-type="doi">10.1007/s10278-023-00776-2</pub-id><pub-id pub-id-type="pmid">36690775</pub-id></citation>
</ref>
<ref id="B103">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Qiao</surname> <given-names>N.</given-names></name> <name><surname>Mostafa</surname> <given-names>H.</given-names></name> <name><surname>Adi</surname> <given-names>F.</given-names></name> <name><surname>Osswald</surname> <given-names>M.</given-names></name> <name><surname>Stefanini</surname> <given-names>F.</given-names></name> <name><surname>Sumislawska</surname> <given-names>D.</given-names></name> <etal/></person-group>. (<year>2015</year>). <article-title>A reconfigurable on-line learning spiking neuromorphic processor comprising 256 neurons and 128k synapses</article-title>. <source>Front. Neurosci</source>. <volume>9</volume>:<fpage>141</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2015.00141</pub-id><pub-id pub-id-type="pmid">25972778</pub-id></citation>
</ref>
<ref id="B104">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Qin</surname> <given-names>L.</given-names></name> <name><surname>Yan</surname> <given-names>R.</given-names></name> <name><surname>Tang</surname> <given-names>H.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;A low latency adaptive coding spike framework for deep reinforcement learning,&#x0201D;</article-title> in <source>Proceedings of the Thirty-Second International Joint Conference on Artificial Intelligence (IJCAI)</source> (<publisher-loc>Macao</publisher-loc>), <fpage>3049</fpage>&#x02013;<lpage>3057</lpage>.</citation>
</ref>
<ref id="B105">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ran&#x000E7;on</surname> <given-names>U.</given-names></name> <name><surname>Cuadrado-Anibarro</surname> <given-names>J.</given-names></name> <name><surname>Cottereau</surname> <given-names>B. R.</given-names></name> <name><surname>Masquelier</surname> <given-names>T.</given-names></name></person-group> (<year>2022</year>). <article-title>Stereospike: depth learning with a spiking neural network</article-title>. <source>IEEE Access</source> <volume>10</volume>, <fpage>127428</fpage>&#x02013;<lpage>127439</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2022.3226484</pub-id></citation>
</ref>
<ref id="B106">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rasmussen</surname> <given-names>D.</given-names></name></person-group> (<year>2019</year>). <article-title>NengoDL: combining deep learning and neuromorphic modelling methods</article-title>. <source>Neuroinformatics</source> <volume>17</volume>, <fpage>611</fpage>&#x02013;<lpage>628</lpage>. <pub-id pub-id-type="doi">10.1007/s12021-019-09424-z</pub-id><pub-id pub-id-type="pmid">30972529</pub-id></citation>
</ref>
<ref id="B107">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rathi</surname> <given-names>N.</given-names></name> <name><surname>Roy</surname> <given-names>K.</given-names></name></person-group> (<year>2023</year>). <article-title>DIET-SNN: a low-latency spiking neural network with direct input encoding and leakage and threshold optimization</article-title>. <source>IEEE Transact. Neural Netw. Learn. Syst</source>. <volume>34</volume>, <fpage>3174</fpage>&#x02013;<lpage>3182</lpage>. <pub-id pub-id-type="doi">10.1109/TNNLS.2021.3111897</pub-id></citation>
</ref>
<ref id="B108">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Rathi</surname> <given-names>N.</given-names></name> <name><surname>Srinivasan</surname> <given-names>G.</given-names></name> <name><surname>Panda</surname> <given-names>P.</given-names></name> <name><surname>Roy</surname> <given-names>K.</given-names></name></person-group> (<year>2020</year>). <article-title>&#x0201C;Enabling deep spiking neural networks with hybrid conversion and spike timing dependent backpropagation,&#x0201D;</article-title> in <source>International Conference on Learning Representations (ICLR)</source> (<publisher-loc>Addis Ababa</publisher-loc>).</citation>
</ref>
<ref id="B109">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ren</surname> <given-names>D.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Peng</surname> <given-names>W.</given-names></name> <name><surname>Liu</surname> <given-names>X.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>&#x0201C;Spiking pointnet: spiking neural networks for point clouds,&#x0201D;</article-title> in <source>Advances in Neural Information Processing Systems (NeurIPS</source>), <italic>Vol. 36</italic> (New Orleans, LA).</citation>
</ref>
<ref id="B110">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ren</surname> <given-names>H.</given-names></name> <name><surname>Zhou</surname> <given-names>Y.</given-names></name> <name><surname>Huang</surname> <given-names>Y.</given-names></name> <name><surname>Fu</surname> <given-names>H.</given-names></name> <name><surname>Lin</surname> <given-names>X.</given-names></name> <name><surname>Song</surname> <given-names>J.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>SpikePoint: an efficient point-based spiking neural network for event cameras action recognition</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2310.07189</pub-id></citation>
</ref>
<ref id="B111">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Roy</surname> <given-names>K.</given-names></name> <name><surname>Jaiswal</surname> <given-names>A.</given-names></name> <name><surname>Panda</surname> <given-names>P.</given-names></name></person-group> (<year>2019</year>). <article-title>Towards spike-based machine intelligence with neuromorphic computing</article-title>. <source>Nature</source> <volume>575</volume>, <fpage>607</fpage>&#x02013;<lpage>617</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-019-1677-2</pub-id><pub-id pub-id-type="pmid">31776490</pub-id></citation>
</ref>
<ref id="B112">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rueckauer</surname> <given-names>B.</given-names></name> <name><surname>Lungu</surname> <given-names>I.-A.</given-names></name> <name><surname>Hu</surname> <given-names>Y.</given-names></name> <name><surname>Pfeiffer</surname> <given-names>M.</given-names></name> <name><surname>Liu</surname> <given-names>S.-C.</given-names></name></person-group> (<year>2017</year>). <article-title>Conversion of continuous-valued deep networks to efficient event-driven networks for image classification</article-title>. <source>Front. Neurosci</source>. <volume>11</volume>:<fpage>682</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2017.00682</pub-id><pub-id pub-id-type="pmid">29375284</pub-id></citation>
</ref>
<ref id="B113">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Schemmel</surname> <given-names>J.</given-names></name> <name><surname>Br&#x000FC;derle</surname> <given-names>D.</given-names></name> <name><surname>Gr&#x000FC;bl</surname> <given-names>A.</given-names></name> <name><surname>Hock</surname> <given-names>M.</given-names></name> <name><surname>Meier</surname> <given-names>K.</given-names></name> <name><surname>Millner</surname> <given-names>S.</given-names></name></person-group> (<year>2010</year>). <article-title>&#x0201C;A wafer-scale neuromorphic hardware system for large-scale neural modeling,&#x0201D;</article-title> in <source>International Symposium on Circuits and Systems (ISCAS)</source> (<publisher-loc>Paris</publisher-loc>), <fpage>1947</fpage>&#x02013;<lpage>1950</lpage>.</citation>
</ref>
<ref id="B114">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Schuman</surname> <given-names>C.</given-names></name> <name><surname>Plank</surname> <given-names>J.</given-names></name> <name><surname>Rose</surname> <given-names>G.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Application-hardware co-design: System-level optimization of neuromorphic computers with neuromorphic devices,&#x0201D;</article-title> in <source>2022 International Electron Devices Meeting (IEDM)</source> (<publisher-loc>San Francisco, CA</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>2</fpage>&#x02013;<lpage>4</lpage>.</citation>
</ref>
<ref id="B115">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Shan</surname> <given-names>L.</given-names></name> <name><surname>Hu</surname> <given-names>B.</given-names></name> <name><surname>Chen</surname> <given-names>L.</given-names></name> <name><surname>Guan</surname> <given-names>Z.-H.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Detecting covid-19 on CT images with impulsive-backpropagation neural networks,&#x0201D;</article-title> in <source>Chinese Control and Decision Conference (CCDC)</source> (<publisher-loc>Hefei</publisher-loc>), <fpage>2797</fpage>&#x02013;<lpage>2803</lpage>.</citation>
</ref>
<ref id="B116">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shen</surname> <given-names>G.</given-names></name> <name><surname>Zhao</surname> <given-names>D.</given-names></name> <name><surname>Zeng</surname> <given-names>Y.</given-names></name></person-group> (<year>2023</year>). <article-title>Exploiting high performance spiking neural networks with efficient spiking patterns</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2301.12356</pub-id></citation>
</ref>
<ref id="B117">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shen</surname> <given-names>J.</given-names></name> <name><surname>Ma</surname> <given-names>D.</given-names></name> <name><surname>Gu</surname> <given-names>Z.</given-names></name> <name><surname>Zhang</surname> <given-names>M.</given-names></name> <name><surname>Zhu</surname> <given-names>X.</given-names></name> <name><surname>Xu</surname> <given-names>X.</given-names></name> <etal/></person-group>. (<year>2016</year>). <article-title>Darwin: a neuromorphic hardware co-processor based on spiking neural networks</article-title>. <source>Sci. China Inf. Sci</source>. <volume>59</volume>, <fpage>1</fpage>&#x02013;<lpage>5</lpage>. <pub-id pub-id-type="doi">10.1007/s11432-015-5511-7</pub-id></citation>
</ref>
<ref id="B118">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Shi</surname> <given-names>X.</given-names></name> <name><surname>Hao</surname> <given-names>Z.</given-names></name> <name><surname>Yu</surname> <given-names>Z.</given-names></name></person-group> (<year>2024</year>). <article-title>&#x0201C;SpikingResformer: bridging resnet and vision transformer in spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)</source> (<publisher-loc>Seattle, WA</publisher-loc>), <fpage>5610</fpage>&#x02013;<lpage>5619</lpage>.</citation>
</ref>
<ref id="B119">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Shrestha</surname> <given-names>S. B.</given-names></name> <name><surname>Orchard</surname> <given-names>G.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Slayer: spike layer error reassignment in time,&#x0201D;</article-title> in <source>Advances in Neural Information Processing Systems (NeurIPS), Vol. 31</source> (<publisher-loc>Montr&#x000E9;al, QC</publisher-loc>).</citation>
</ref>
<ref id="B120">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Soures</surname> <given-names>N.</given-names></name> <name><surname>Kudithipudi</surname> <given-names>D.</given-names></name></person-group> (<year>2019</year>). <article-title>Deep liquid state machines with neural plasticity for video activity recognition</article-title>. <source>Front. Neurosci</source>. <volume>13</volume>:<fpage>457929</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2019.00686</pub-id></citation>
</ref>
<ref id="B121">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Su</surname> <given-names>Q.</given-names></name> <name><surname>Chou</surname> <given-names>Y.</given-names></name> <name><surname>Hu</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Mei</surname> <given-names>S.</given-names></name> <name><surname>Zhang</surname> <given-names>Z.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>&#x0201C;Deep directly-trained spiking neural networks for object detection,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)</source> (<publisher-loc>Paris</publisher-loc>), <fpage>6555</fpage>&#x02013;<lpage>6565</lpage>.</citation>
</ref>
<ref id="B122">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Vaswani</surname> <given-names>A.</given-names></name> <name><surname>Shazeer</surname> <given-names>N.</given-names></name> <name><surname>Parmar</surname> <given-names>N.</given-names></name> <name><surname>Uszkoreit</surname> <given-names>J.</given-names></name> <name><surname>Jones</surname> <given-names>L.</given-names></name> <name><surname>Gomez</surname> <given-names>A. N.</given-names></name> <etal/></person-group>. (<year>2017</year>). <article-title>&#x0201C;Attention is all you need,&#x0201D;</article-title> in <source>Advances in Neural Information Processing Systems (NeurIPS), Vol. 30</source> (<publisher-loc>Long Beach, CA</publisher-loc>).</citation>
</ref>
<ref id="B123">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>B.</given-names></name> <name><surname>Dong</surname> <given-names>G.</given-names></name> <name><surname>Zhao</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>R.</given-names></name> <name><surname>Yang</surname> <given-names>H.</given-names></name> <name><surname>Yin</surname> <given-names>W.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>&#x0201C;Spiking emotions: dynamic vision emotion recognition using spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the International Conference on Algorithms, High Performance Computing and Artificial Intelligence (AHPCAI)</source> (<publisher-loc>Guangzhou</publisher-loc>), <fpage>50</fpage>&#x02013;<lpage>58</lpage>.</citation>
</ref>
<ref id="B124">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>H.</given-names></name> <name><surname>Li</surname> <given-names>Y.-F.</given-names></name></person-group> (<year>2023</year>). <article-title>Bioinspired membrane learnable spiking neural network for autonomous vehicle sensors fault diagnosis under open environments</article-title>. <source>Reliabil. Eng. Syst. Saf</source>. <volume>233</volume>:<fpage>109102</fpage>. <pub-id pub-id-type="doi">10.1016/j.ress.2023.109102</pub-id></citation>
</ref>
<ref id="B125">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>H.</given-names></name> <name><surname>Sun</surname> <given-names>M.</given-names></name> <name><surname>Li</surname> <given-names>Y.-F.</given-names></name></person-group> (<year>2023b</year>). <article-title>A brain-inspired spiking network framework based on multi-time-step self-attention for lithium-ion batteries capacity prediction</article-title>. <source>IEEE Transact. Consum. Electron</source>. <volume>70</volume>:<fpage>3279135</fpage>. <pub-id pub-id-type="doi">10.1109/TCE.2023.3279135</pub-id></citation>
</ref>
<ref id="B126">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>M.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Ma</surname> <given-names>M.</given-names></name> <name><surname>Fan</surname> <given-names>X.</given-names></name></person-group> (<year>2023</year>). <article-title>Spiking semantic communication for feature transmission with Harq</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2310.08804</pub-id></citation>
</ref>
<ref id="B127">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>Q.</given-names></name> <name><surname>Li</surname> <given-names>P.</given-names></name></person-group> (<year>2016</year>). <article-title>&#x0201C;D-LSM: deep liquid state machine with unsupervised recurrent reservoir tuning,&#x0201D;</article-title> in <source>2016 23rd International Conference on Pattern Recognition (ICPR)</source> (<publisher-loc>IEEE</publisher-loc>), <fpage>2652</fpage>&#x02013;<lpage>2657</lpage>.</citation>
</ref>
<ref id="B128">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>R.</given-names></name> <name><surname>Thakur</surname> <given-names>C. S.</given-names></name> <name><surname>Cohen</surname> <given-names>G.</given-names></name> <name><surname>Hamilton</surname> <given-names>T. J.</given-names></name> <name><surname>Tapson</surname> <given-names>J.</given-names></name> <name><surname>van Schaik</surname> <given-names>A.</given-names></name></person-group> (<year>2017</year>). <article-title>Neuromorphic hardware architecture using the neural engineering framework for pattern recognition</article-title>. <source>IEEE Trans. Biomed. Circuits Syst</source>. <volume>11</volume>, <fpage>574</fpage>&#x02013;<lpage>584</lpage>. <pub-id pub-id-type="doi">10.1109/TBCAS.2017.2666883</pub-id><pub-id pub-id-type="pmid">28436888</pub-id></citation>
</ref>
<ref id="B129">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>S.</given-names></name> <name><surname>Cheng</surname> <given-names>T. H.</given-names></name> <name><surname>Lim</surname> <given-names>M. H.</given-names></name></person-group> (<year>2022</year>). <article-title>LTMD: learning improvement of spiking neural networks with learnable thresholding neurons and moderate dropout</article-title>. <source>Adv. Neural Inf. Process. Syst</source>. <volume>35</volume>, <fpage>28350</fpage>&#x02013;<lpage>28362</lpage>.</citation>
</ref>
<ref id="B130">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>W.</given-names></name> <name><surname>Hao</surname> <given-names>S.</given-names></name> <name><surname>Wei</surname> <given-names>Y.</given-names></name> <name><surname>Xiao</surname> <given-names>S.</given-names></name> <name><surname>Feng</surname> <given-names>J.</given-names></name> <name><surname>Sebe</surname> <given-names>N.</given-names></name></person-group> (<year>2019</year>). <article-title>Temporal spiking recurrent neural network for action recognition</article-title>. <source>IEEE Access</source> <volume>7</volume>, <fpage>117165</fpage>&#x02013;<lpage>117175</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2019.2936604</pub-id></citation>
</ref>
<ref id="B131">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>W.</given-names></name> <name><surname>Xie</surname> <given-names>E.</given-names></name> <name><surname>Li</surname> <given-names>X.</given-names></name> <name><surname>Fan</surname> <given-names>D.-P.</given-names></name> <name><surname>Song</surname> <given-names>K.</given-names></name> <name><surname>Liang</surname> <given-names>D.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>&#x0201C;Pyramid vision transformer: a versatile backbone for dense prediction without convolutions,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)</source>, <fpage>568</fpage>&#x02013;<lpage>578</lpage>.</citation>
</ref>
<ref id="B132">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>X.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name></person-group> (<year>2023</year>). <article-title>MT-SNN: enhance spiking neural network with multiple thresholds</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2303.11127</pub-id></citation>
</ref>
<ref id="B133">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>H.</given-names></name> <name><surname>Li</surname> <given-names>Y.-F.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name></person-group> (<year>2023a</year>). <article-title>Bioinspired spiking spatiotemporal attention framework for lithium-ion batteries state-of-health estimation</article-title>. <source>Renew. Sustain. Energy Rev</source>. <volume>188</volume>:<fpage>113728</fpage>. <pub-id pub-id-type="doi">10.1016/j.rser.2023.113728</pub-id></citation>
</ref>
<ref id="B134">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Shi</surname> <given-names>K.</given-names></name> <name><surname>Lu</surname> <given-names>C.</given-names></name> <name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>M.</given-names></name> <name><surname>Qu</surname> <given-names>H.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;Spatial-temporal self-attention for asynchronous spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the Thirty-Second International Joint Conference on Artificial Intelligence (IJCAI)</source> (<publisher-loc>Macao SAR</publisher-loc>), <fpage>3085</fpage>&#x02013;<lpage>3093</lpage>.</citation>
</ref>
<ref id="B135">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>M.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Qu</surname> <given-names>H.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Signed neuron with memory: Towards simple, accurate and high-efficient ann-snn conversion,&#x0201D;</article-title> in <source>Proceedings of the Thirty-First International Joint Conference on Artificial Intelligence (IJCAI)</source> (<publisher-loc>Vienna</publisher-loc>), <fpage>2501</fpage>&#x02013;<lpage>2508</lpage>.</citation>
</ref>
<ref id="B136">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Werbos</surname> <given-names>P.</given-names></name></person-group> (<year>1990</year>). <article-title>Backpropagation through time: what it does and how to do it</article-title>. <source>Proc. IEEE</source> <volume>78</volume>, <fpage>1550</fpage>&#x02013;<lpage>1560</lpage>. <pub-id pub-id-type="doi">10.1109/5.58337</pub-id></citation>
</ref>
<ref id="B137">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wiener</surname> <given-names>M. C.</given-names></name> <name><surname>Richmond</surname> <given-names>B. J.</given-names></name></person-group> (<year>2003</year>). <article-title>Decoding spike trains instant by instant using order statistics and the mixture-of-poissons model</article-title>. <source>J. Neurosci</source>. <volume>23</volume>, <fpage>2394</fpage>&#x02013;<lpage>2406</lpage>. <pub-id pub-id-type="doi">10.1523/JNEUROSCI.23-06-02394.2003</pub-id><pub-id pub-id-type="pmid">12657699</pub-id></citation>
</ref>
<ref id="B138">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wu</surname> <given-names>M.-H.</given-names></name> <name><surname>Huang</surname> <given-names>M.-S.</given-names></name> <name><surname>Zhu</surname> <given-names>Z.</given-names></name> <name><surname>Liang</surname> <given-names>F.-X.</given-names></name> <name><surname>Hong</surname> <given-names>M.-C.</given-names></name> <name><surname>Deng</surname> <given-names>J.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>&#x0201C;Compact probabilistic poisson neuron based on back-hopping oscillation in stt-mram for all-spin deep spiking neural network,&#x0201D;</article-title> in <source>Symposium on VLSI Technology</source>, <fpage>1</fpage>&#x02013;<lpage>2</lpage>.</citation>
</ref>
<ref id="B139">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wu</surname> <given-names>X.</given-names></name> <name><surname>He</surname> <given-names>W.</given-names></name> <name><surname>Yao</surname> <given-names>M.</given-names></name> <name><surname>Zhang</surname> <given-names>Z.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>G.</given-names></name></person-group> (<year>2022</year>). <article-title>MSS-DepthNet: depth prediction with multi-step spiking neural network</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2211.12156</pub-id></citation>
</ref>
<ref id="B140">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wu</surname> <given-names>Y.</given-names></name> <name><surname>Deng</surname> <given-names>L.</given-names></name> <name><surname>Li</surname> <given-names>G.</given-names></name> <name><surname>Zhu</surname> <given-names>J.</given-names></name> <name><surname>Shi</surname> <given-names>L.</given-names></name></person-group> (<year>2018</year>). <article-title>Spatio-temporal backpropagation for training high-performance spiking neural networks</article-title>. <source>Front. Neurosci</source>. <volume>12</volume>:<fpage>331</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2018.00331</pub-id><pub-id pub-id-type="pmid">29875621</pub-id></citation>
</ref>
<ref id="B141">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wu</surname> <given-names>Y.</given-names></name> <name><surname>Deng</surname> <given-names>L.</given-names></name> <name><surname>Li</surname> <given-names>G.</given-names></name> <name><surname>Zhu</surname> <given-names>J.</given-names></name> <name><surname>Xie</surname> <given-names>Y.</given-names></name> <name><surname>Shi</surname> <given-names>L.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Direct training for spiking neural networks: Faster, larger, better,&#x0201D;</article-title> in <source>Proceedings of the AAAI Conference on Artificial Intelligence (AAAI)</source>, <fpage>1311</fpage>&#x02013;<lpage>1318</lpage>.</citation>
</ref>
<ref id="B142">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Xiang</surname> <given-names>S.</given-names></name> <name><surname>Zhang</surname> <given-names>T.</given-names></name> <name><surname>Jiang</surname> <given-names>S.</given-names></name> <name><surname>Han</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Du</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Spiking SiamFC&#x0002B;&#x0002B;: deep spiking neural network for object tracking</article-title>. <source>arXiv [Preprint]</source>. abs/2209.12010.</citation>
</ref>
<ref id="B143">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xiao</surname> <given-names>M.</given-names></name> <name><surname>Meng</surname> <given-names>Q.</given-names></name> <name><surname>Zhang</surname> <given-names>Z.</given-names></name> <name><surname>He</surname> <given-names>D.</given-names></name> <name><surname>Lin</surname> <given-names>Z.</given-names></name></person-group> (<year>2022</year>). <article-title>Online training through time for spiking neural networks</article-title>. <source>Adv. Neural Inf. Process. Syst</source>. <volume>35</volume>, <fpage>20717</fpage>&#x02013;<lpage>20730</lpage>.</citation>
</ref>
<ref id="B144">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Xing</surname> <given-names>X.</given-names></name> <name><surname>Zhang</surname> <given-names>Z.</given-names></name> <name><surname>Ni</surname> <given-names>Z.</given-names></name> <name><surname>Xiao</surname> <given-names>S.</given-names></name> <name><surname>Ju</surname> <given-names>Y.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>&#x0201C;SpikeLM: towards general spike-driven language modeling via elastic bi-spiking mechanisms,&#x0201D;</article-title> in <source>Forty-First International Conference on Machine Learning (ICML)</source> (<publisher-loc>Vienna</publisher-loc>).</citation>
</ref>
<ref id="B145">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xiong</surname> <given-names>Y.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Chen</surname> <given-names>C.</given-names></name> <name><surname>Wei</surname> <given-names>X.</given-names></name> <name><surname>Xue</surname> <given-names>Y.</given-names></name> <name><surname>Wan</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>An odor recognition algorithm of electronic noses based on convolutional spiking neural network for spoiled food identification</article-title>. <source>J. Electrochem. Soc</source>. <volume>168</volume>:<fpage>077519</fpage>. <pub-id pub-id-type="doi">10.1149/1945-7111/ac1699</pub-id></citation>
</ref>
<ref id="B146">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xu</surname> <given-names>Z.</given-names></name> <name><surname>Ma</surname> <given-names>Y.</given-names></name> <name><surname>Pan</surname> <given-names>Z.</given-names></name> <name><surname>Zheng</surname> <given-names>X.</given-names></name></person-group> (<year>2022</year>). <article-title>Deep spiking residual shrinkage network for bearing fault diagnosis</article-title>. <source>IEEE Trans. Cybern</source>. <volume>54</volume>, <fpage>1</fpage>&#x02013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1109/TCYB.2022.3227363</pub-id></citation>
</ref>
<ref id="B147">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yamazaki</surname> <given-names>K.</given-names></name> <name><surname>Vo-Ho</surname> <given-names>V.-K.</given-names></name> <name><surname>Bulsara</surname> <given-names>D.</given-names></name> <name><surname>Le</surname> <given-names>N.</given-names></name></person-group> (<year>2022</year>). <article-title>Spiking neural networks and their applications: a review</article-title>. <source>Brain Sci</source>. <volume>12</volume>:<fpage>863</fpage>. <pub-id pub-id-type="doi">10.3390/brainsci12070863</pub-id><pub-id pub-id-type="pmid">35884670</pub-id></citation>
</ref>
<ref id="B148">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yang</surname> <given-names>Z.</given-names></name> <name><surname>Wu</surname> <given-names>Y.</given-names></name> <name><surname>Wang</surname> <given-names>G.</given-names></name> <name><surname>Yang</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>G.</given-names></name> <name><surname>Deng</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2019</year>). <article-title>DashNet: a hybrid artificial and spiking neural network for high-speed object tracking</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.1909.12942</pub-id></citation>
</ref>
<ref id="B149">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yao</surname> <given-names>M.</given-names></name> <name><surname>Hu</surname> <given-names>J.</given-names></name> <name><surname>Hu</surname> <given-names>T.</given-names></name> <name><surname>Xu</surname> <given-names>Y.</given-names></name> <name><surname>Zhou</surname> <given-names>Z.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>&#x0201C;Spike-driven transformer v2: meta spiking neural network architecture inspiring the design of next-generation neuromorphic chips,&#x0201D;</article-title> in <source>The Twelfth International Conference on Learning Representations (ICLR)</source>.</citation>
</ref>
<ref id="B150">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yao</surname> <given-names>M.</given-names></name> <name><surname>Zhao</surname> <given-names>G.</given-names></name> <name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Hu</surname> <given-names>Y.</given-names></name> <name><surname>Deng</surname> <given-names>L.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2023b</year>). <article-title>Attention spiking neural networks</article-title>. <source>IEEE Trans. Pattern Anal. Mach. Intell</source>. <volume>45</volume>, <fpage>9393</fpage>&#x02013;<lpage>9410</lpage>. <pub-id pub-id-type="doi">10.1109/TPAMI.2023.3241201</pub-id><pub-id pub-id-type="pmid">37022261</pub-id></citation>
</ref>
<ref id="B151">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yao</surname> <given-names>X.</given-names></name> <name><surname>Li</surname> <given-names>F.</given-names></name> <name><surname>Mo</surname> <given-names>Z.</given-names></name> <name><surname>Cheng</surname> <given-names>J.</given-names></name></person-group> (<year>2022</year>). <article-title>GLIF: a unified gated leaky integrate-and-fire neuron for spiking neural networks</article-title>. <source>Adv. Neural Inf. Process. Syst</source>. <volume>35</volume>, <fpage>32160</fpage>&#x02013;<lpage>32171</lpage>.</citation>
</ref>
<ref id="B152">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Yao</surname> <given-names>M.</given-names></name> <name><surname>Hu</surname> <given-names>J.</given-names></name> <name><surname>Zhou</surname> <given-names>Z.</given-names></name> <name><surname>Yuan</surname> <given-names>L.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name> <name><surname>Xu</surname> <given-names>B.</given-names></name> <etal/></person-group>. (<year>2023a</year>). <article-title>&#x0201C;Spike-driven transformer,&#x0201D;</article-title> in <source>Advances in Neural Information Processing Systems (NeurIPS), Vol. 36</source> (<publisher-loc>New Orleans, LA</publisher-loc>).</citation>
</ref>
<ref id="B153">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Yarga</surname> <given-names>S. Y. A.</given-names></name> <name><surname>Wood</surname> <given-names>S. U. N.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201C;Accelerating SNN training with stochastic parallelizable spiking neurons,&#x0201D;</article-title> in <source>International Joint Conference on Neural Networks (IJCNN)</source> (<publisher-loc>Gold Coast, QLD</publisher-loc>), <fpage>1</fpage>&#x02013;<lpage>8</lpage>.</citation>
</ref>
<ref id="B154">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Yin</surname> <given-names>N.</given-names></name> <name><surname>Wang</surname> <given-names>M.</given-names></name> <name><surname>Chen</surname> <given-names>Z.</given-names></name> <name><surname>De Masi</surname> <given-names>G.</given-names></name> <name><surname>Xiong</surname> <given-names>H.</given-names></name> <name><surname>Gu</surname> <given-names>B.</given-names></name></person-group> (<year>2024</year>). <article-title>&#x0201C;Dynamic spiking graph neural networks,&#x0201D;</article-title> in <source>Proceedings of the AAAI Conference on Artificial Intelligence (AAAI)</source> (<publisher-loc>Vancouver, BC</publisher-loc>), <fpage>16495</fpage>&#x02013;<lpage>16503</lpage>.</citation>
</ref>
<ref id="B155">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yin</surname> <given-names>R.</given-names></name> <name><surname>Moitra</surname> <given-names>A.</given-names></name> <name><surname>Bhattacharjee</surname> <given-names>A.</given-names></name> <name><surname>Kim</surname> <given-names>Y.</given-names></name> <name><surname>Panda</surname> <given-names>P.</given-names></name></person-group> (<year>2022</year>). <article-title>SATA: sparsity-aware training accelerator for spiking neural networks</article-title>. <source>IEEE Transact. Comp. Aided Des. Integr. Circ. Syst</source>. 42. 1926&#x02013;1938. <pub-id pub-id-type="doi">10.1109/TCAD.2022.3213211</pub-id></citation>
</ref>
<ref id="B156">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yin</surname> <given-names>B.</given-names></name> <name><surname>Corradi</surname> <given-names>F.</given-names></name> <name><surname>Boht&#x000E9;</surname> <given-names>S. M.</given-names></name></person-group> (<year>2020</year>). <article-title>Effective and efficient computation with multiple-timescale spiking recurrent neural networks</article-title>. <source>Int. Conf. Neuromor. Syst</source>. <volume>2020</volume>, <fpage>1</fpage>&#x02013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1145/3407197.3407225</pub-id></citation>
</ref>
<ref id="B157">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yu</surname> <given-names>L.</given-names></name> <name><surname>Huang</surname> <given-names>L.</given-names></name> <name><surname>Zhou</surname> <given-names>C.</given-names></name> <name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Zhou</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>SVFormer: a direct training spiking transformer for efficient video action recognition</article-title>. <source>arXiv [Preprint]</source>. abs/2406.15034. <pub-id pub-id-type="doi">10.48550/arXiv.2406.15034</pub-id></citation>
</ref>
<ref id="B158">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yuan</surname> <given-names>L.</given-names></name> <name><surname>Hou</surname> <given-names>Q.</given-names></name> <name><surname>Jiang</surname> <given-names>Z.</given-names></name> <name><surname>Feng</surname> <given-names>J.</given-names></name> <name><surname>Yan</surname> <given-names>S.</given-names></name></person-group> (<year>2022</year>). <article-title>VOLO: vision outlooker for visual recognition</article-title>. <source>IEEE Trans. Pattern Anal. Mach. Intell</source>. <volume>45</volume>, <fpage>6575</fpage>&#x02013;<lpage>6586</lpage>. <pub-id pub-id-type="doi">10.1109/TPAMI.2022.3206108</pub-id></citation>
</ref>
<ref id="B159">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Yuan</surname> <given-names>L.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Wang</surname> <given-names>T.</given-names></name> <name><surname>Yu</surname> <given-names>W.</given-names></name> <name><surname>Shi</surname> <given-names>Y.</given-names></name> <name><surname>Jiang</surname> <given-names>Z.-H.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>&#x0201C;Tokens-to-token vit: training vision transformers from scratch on imagenet,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)</source> (<publisher-loc>Montreal, QC</publisher-loc>), <fpage>558</fpage>&#x02013;<lpage>567</lpage>.</citation>
</ref>
<ref id="B160">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Fan</surname> <given-names>X.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name></person-group> (<year>2023</year>). <article-title>Energy-efficient spiking segmenter for frame and event-based images</article-title>. <source>Biomimetics</source> <volume>8</volume>:<fpage>356</fpage>. <pub-id pub-id-type="doi">10.3390/biomimetics8040356</pub-id><pub-id pub-id-type="pmid">37622961</pub-id></citation>
</ref>
<ref id="B161">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Zhou</surname> <given-names>C.</given-names></name> <name><surname>Yu</surname> <given-names>L.</given-names></name> <name><surname>Huang</surname> <given-names>L.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Fan</surname> <given-names>X.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>SGLFormer: spiking global-local-fusion transformer with high performance</article-title>. <source>Front. Neurosci</source>. <volume>18</volume>:<fpage>1371290</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2024.1371290</pub-id><pub-id pub-id-type="pmid">38550564</pub-id></citation>
</ref>
<ref id="B162">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>J.</given-names></name> <name><surname>Tang</surname> <given-names>L.</given-names></name> <name><surname>Yu</surname> <given-names>Z.</given-names></name> <name><surname>Lu</surname> <given-names>J.</given-names></name> <name><surname>Huang</surname> <given-names>T.</given-names></name></person-group> (<year>2022b</year>). <article-title>Spike transformer: monocular depth estimation for spiking camera</article-title>. <source>Proc. Eur. Conf. Comp. Vis</source>. <volume>13667</volume>, <fpage>34</fpage>&#x02013;<lpage>52</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-031-20071-7_3</pub-id></citation>
</ref>
<ref id="B163">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>J.</given-names></name> <name><surname>Wang</surname> <given-names>J.</given-names></name> <name><surname>Di</surname> <given-names>X.</given-names></name> <name><surname>Pu</surname> <given-names>S.</given-names></name></person-group> (<year>2022c</year>). <article-title>&#x0201C;High-accuracy and energy-efficient action recognition with deep spiking neural network,&#x0201D;</article-title> in <source>International Conference on Neural Information Processing (ICONIP)</source>, <fpage>279</fpage>&#x02013;<lpage>292</lpage>.</citation>
</ref>
<ref id="B164">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>X.</given-names></name> <name><surname>Lu</surname> <given-names>J.</given-names></name> <name><surname>Wang</surname> <given-names>Z.</given-names></name> <name><surname>Wang</surname> <given-names>R.</given-names></name> <name><surname>Wei</surname> <given-names>J.</given-names></name> <name><surname>Shi</surname> <given-names>T.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Hybrid memristor-cmos neurons for <italic>in-situ</italic> learning in fully hardware memristive spiking neural networks</article-title>. <source>Sci. Bull</source>. <volume>66</volume>, <fpage>1624</fpage>&#x02013;<lpage>1633</lpage>. <pub-id pub-id-type="doi">10.1016/j.scib.2021.04.014</pub-id></citation>
</ref>
<ref id="B165">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>Z.</given-names></name> <name><surname>Lew</surname> <given-names>L.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;PokeBNN: a binary pursuit of lightweight accuracy,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)</source>, <fpage>12475</fpage>&#x02013;<lpage>12485</lpage>.</citation>
</ref>
<ref id="B166">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>J.</given-names></name> <name><surname>Huang</surname> <given-names>L.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Zhou</surname> <given-names>H.</given-names></name></person-group> (<year>2023</year>). <article-title>Predicting the temporal-dynamic trajectories of cortical neuronal responses in non-human primates based on deep spiking neural network</article-title>. <source>Cogn. Neurodyn</source>. 1&#x02013;12. <pub-id pub-id-type="doi">10.1007/s11571-023-09989-1</pub-id></citation>
</ref>
<ref id="B167">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>J.</given-names></name> <name><surname>Dong</surname> <given-names>B.</given-names></name> <name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Ding</surname> <given-names>J.</given-names></name> <name><surname>Heide</surname> <given-names>F.</given-names></name> <name><surname>Yin</surname> <given-names>B.</given-names></name> <etal/></person-group>. (<year>2022a</year>). <article-title>&#x0201C;Spiking transformers for event-based single object tracking,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)</source> (<publisher-loc>New Orleans, LA</publisher-loc>), <fpage>8801</fpage>&#x02013;<lpage>8810</lpage>.</citation>
</ref>
<ref id="B168">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zheng</surname> <given-names>H.</given-names></name> <name><surname>Wu</surname> <given-names>Y.</given-names></name> <name><surname>Deng</surname> <given-names>L.</given-names></name> <name><surname>Hu</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>G.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;Going deeper with directly-trained larger spiking neural networks,&#x0201D;</article-title> in <source>Proceedings of the AAAI Conference on Artificial Intelligence (AAAI)</source> (<publisher-loc>Vancouver, BC</publisher-loc>), <fpage>11062</fpage>&#x02013;<lpage>11070</lpage>.</citation>
</ref>
<ref id="B169">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhou</surname> <given-names>C.</given-names></name> <name><surname>Yu</surname> <given-names>L.</given-names></name> <name><surname>Zhou</surname> <given-names>Z.</given-names></name> <name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Zhou</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2023a</year>). <article-title>Spikingformer: spike-driven residual learning for transformer-based spiking neural network</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2304.11954</pub-id></citation>
</ref>
<ref id="B170">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhou</surname> <given-names>C.</given-names></name> <name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Zhou</surname> <given-names>Z.</given-names></name> <name><surname>Yu</surname> <given-names>L.</given-names></name> <name><surname>Huang</surname> <given-names>L.</given-names></name> <name><surname>Fan</surname> <given-names>X.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>QKFormer: hierarchical spiking transformer using qk attention</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2403.16552</pub-id></citation>
</ref>
<ref id="B171">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhou</surname> <given-names>C.</given-names></name> <name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Zhou</surname> <given-names>Z.</given-names></name> <name><surname>Yu</surname> <given-names>L.</given-names></name> <name><surname>Ma</surname> <given-names>Z.</given-names></name> <name><surname>Zhou</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2023b</year>). <article-title>Enhancing the performance of transformer-based spiking neural networks by SNN-optimized downsampling with precise gradient backpropagation</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2305.05954</pub-id></citation>
</ref>
<ref id="B172">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhou</surname> <given-names>S.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>X.</given-names></name> <name><surname>Sanyal</surname> <given-names>A.</given-names></name></person-group> (<year>2020</year>). <article-title>Deep SCNN-based real-time object detection for self-driving vehicles using lidar temporal data</article-title>. <source>IEEE Access</source> <volume>8</volume>, <fpage>76903</fpage>&#x02013;<lpage>76912</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2020.2990416</pub-id></citation>
</ref>
<ref id="B173">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhou</surname> <given-names>Z.</given-names></name> <name><surname>Che</surname> <given-names>K.</given-names></name> <name><surname>Fang</surname> <given-names>W.</given-names></name> <name><surname>Tian</surname> <given-names>K.</given-names></name> <name><surname>Zhu</surname> <given-names>Y.</given-names></name> <name><surname>Yan</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>Spikformer v2: join the high accuracy club on imagenet with an SNN ticket</article-title>. <source>arXiv [Preprint]</source>.</citation>
</ref>
<ref id="B174">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhou</surname> <given-names>Z.</given-names></name> <name><surname>Zhu</surname> <given-names>Y.</given-names></name> <name><surname>He</surname> <given-names>C.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>YAN</surname> <given-names>S.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>&#x0201C;Spikformer: when spiking neural network meets transformer,&#x0201D;</article-title> in <source>International Conference on Learning Representations (ICLR)</source> (<publisher-loc>Kigali</publisher-loc>).</citation>
</ref>
<ref id="B175">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhu</surname> <given-names>L.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Wang</surname> <given-names>X.</given-names></name> <name><surname>Huang</surname> <given-names>T.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;Neuspike-net: High speed video reconstruction via bio-inspired neuromorphic cameras,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF International Conference on Computer Vision (ICCV)</source> (<publisher-loc>Montreal, QC</publisher-loc>), <fpage>2400</fpage>&#x02013;<lpage>2409</lpage>.</citation>
</ref>
<ref id="B176">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhu</surname> <given-names>L.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2023</year>). <article-title>Review of visual reconstruction methods of retina-like vision sensors</article-title>. <source>Sci. Sinica</source> <volume>53</volume>, <fpage>417</fpage>&#x02013;<lpage>436</lpage>. <pub-id pub-id-type="doi">10.1360/SSI-2021-0397</pub-id></citation>
</ref>
<ref id="B177">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhu</surname> <given-names>L.</given-names></name> <name><surname>Wang</surname> <given-names>X.</given-names></name> <name><surname>Chang</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Huang</surname> <given-names>T.</given-names></name> <name><surname>Tian</surname> <given-names>Y.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Event-based video reconstruction via potential-assisted spiking neural network,&#x0201D;</article-title> in <source>Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)</source> (<publisher-loc>New Orleans, LA</publisher-loc>), <fpage>3594</fpage>&#x02013;<lpage>3604</lpage>.</citation>
</ref>
<ref id="B178">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhu</surname> <given-names>R.</given-names></name> <name><surname>Zhao</surname> <given-names>Q.</given-names></name> <name><surname>Eshraghian</surname> <given-names>J. K.</given-names></name></person-group> (<year>2023</year>). <article-title>SpikeGPT: generative pre-trained language model with spiking neural networks</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2302.13939</pub-id></citation>
</ref>
<ref id="B179">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zhu</surname> <given-names>X.</given-names></name> <name><surname>Su</surname> <given-names>W.</given-names></name> <name><surname>Lu</surname> <given-names>L.</given-names></name> <name><surname>Li</surname> <given-names>B.</given-names></name> <name><surname>Wang</surname> <given-names>X.</given-names></name> <name><surname>Dai</surname> <given-names>J.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;Deformable DETR: deformable transformers for end-to-end object detection,&#x0201D;</article-title> in <source>International Conference on Learning Representations (ICLR)</source> (<publisher-loc>Vienna</publisher-loc>).</citation>
</ref>
<ref id="B180">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhu</surname> <given-names>Z.</given-names></name> <name><surname>Peng</surname> <given-names>J.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Chen</surname> <given-names>L.</given-names></name> <name><surname>Yu</surname> <given-names>Q.</given-names></name> <name><surname>Luo</surname> <given-names>S.</given-names></name></person-group> (<year>2022</year>). <article-title>Spiking graph convolutional networks</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.24963/ijcai.2022/338</pub-id></citation>
</ref>
<ref id="B181">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zou</surname> <given-names>S.</given-names></name> <name><surname>Mu</surname> <given-names>Y.</given-names></name> <name><surname>Zuo</surname> <given-names>X.</given-names></name> <name><surname>Wang</surname> <given-names>S.</given-names></name> <name><surname>Li</surname> <given-names>C.</given-names></name></person-group> (<year>2023</year>). <article-title>Event-based human pose tracking by spiking spatiotemporal transformer</article-title>. <source>arXiv [Preprint]</source>. <pub-id pub-id-type="doi">10.48550/arXiv.2303.09681</pub-id></citation>
</ref>
</ref-list>
</back>
</article>