<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="review-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Phys.</journal-id>
<journal-title>Frontiers in Physics</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Phys.</abbrev-journal-title>
<issn pub-type="epub">2296-424X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1369099</article-id>
<article-id pub-id-type="doi">10.3389/fphy.2024.1369099</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Physics</subject>
<subj-group>
<subject>Review</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>A review of emerging trends in photonic deep learning accelerators</article-title>
<alt-title alt-title-type="left-running-head">Atwany et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fphy.2024.1369099">10.3389/fphy.2024.1369099</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Atwany</surname>
<given-names>Mohammad</given-names>
</name>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2627862/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Pardo</surname>
<given-names>Sarah</given-names>
</name>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2711527/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
</contrib>
<contrib contrib-type="author" corresp="yes" equal-contrib="yes">
<name>
<surname>Serunjogi</surname>
<given-names>Solomon</given-names>
</name>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<xref ref-type="author-notes" rid="fn002">
<sup>&#x2021;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2729443/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes" equal-contrib="yes">
<name>
<surname>Rasras</surname>
<given-names>Mahmoud</given-names>
</name>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<xref ref-type="author-notes" rid="fn002">
<sup>&#x2021;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2674380/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff>
<institution>Engineering Division</institution>, <institution>New York University</institution>, <addr-line>Abu Dhabi</addr-line>, <country>United Arab Emirates</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/177473/overview">Supriyo Bandyopadhyay</ext-link>, Virginia Commonwealth University, United States</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1268904/overview">Shuming Jiao</ext-link>, Peng Cheng Laboratory, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/263279/overview">Sunkyu Yu</ext-link>, Seoul National University, Republic of Korea</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Solomon Serunjogi, <email>sms10215@nyu.edu</email>; Mahmoud Rasras, <email>mr5098@nyu.edu</email>
</corresp>
<fn fn-type="equal" id="fn001">
<label>
<sup>&#x2020;</sup>
</label>
<p>These authors have contributed equally to this work and share first authorship</p>
</fn>
<fn fn-type="equal" id="fn002">
<label>
<sup>&#x2021;</sup>
</label>
<p>These authors share senior authorship</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>15</day>
<month>07</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>12</volume>
<elocation-id>1369099</elocation-id>
<history>
<date date-type="received">
<day>11</day>
<month>01</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>17</day>
<month>04</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2024 Atwany, Pardo, Serunjogi and Rasras.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Atwany, Pardo, Serunjogi and Rasras</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Deep learning has revolutionized many sectors of industry and daily life, but as application scale increases, performing training and inference with large models on massive datasets is increasingly unsustainable on existing hardware. Highly parallelized hardware like Graphics Processing Units (GPUs) are now widely used to improve speed over conventional Central Processing Units (CPUs). However, Complementary Metal-oxide Semiconductor (CMOS) devices suffer from fundamental limitations relying on metallic interconnects which impose inherent constraints on bandwidth, latency, and energy efficiency. Indeed, by 2026, the projected global electricity consumption of data centers fueled by CMOS chips is expected to increase by an amount equivalent to the annual usage of an additional European country. Silicon Photonics (SiPh) devices are emerging as a promising energy-efficient CMOS-compatible alternative to electronic deep learning accelerators, using light to compute as well as communicate. In this review, we examine the prospects of photonic computing as an emerging solution for acceleration in deep learning applications. We present an overview of the photonic computing landscape, then focus in detail on SiPh integrated circuit (PIC) accelerators designed for different neural network models and applications deep learning. We categorize different devices based on their use cases and operating principles to assess relative strengths, present open challenges, and identify new directions for further research.</p>
</abstract>
<kwd-group>
<kwd>photonics integrated circuits</kwd>
<kwd>photonic deep learning accelerators</kwd>
<kwd>deep neural networks</kwd>
<kwd>artificial intelligence</kwd>
<kwd>silicon photonics (SiPh)</kwd>
</kwd-group>
<contract-sponsor id="cn001">New York University Abu Dhabi<named-content content-type="fundref-id">10.13039/100012025</named-content>
</contract-sponsor>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Optics and Photonics</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>Since the advent of computers, researchers have been captivated by the prospect of endowing machines with human-like abilities such as abstract thinking, decision-making, creative expression, and social behavior, leading to a field now popularized as <italic>Artificial Intelligence (AI)</italic>. However, the practical use of AI was not realized until theoretical advances in the 1980s and 90s brought a particular form of AI to the forefront: deep neural networks [<xref ref-type="bibr" rid="B1">1</xref>]. In the past two decades, deep learning has seamlessly integrated into daily life, from consumer applications like personalized product recommendations, to enhanced medical diagnosis and drug design. This rapid advancement in adoption can be substantially attributed to advances in hardware, both in processing speed and memory capacity [<xref ref-type="bibr" rid="B2">2</xref>].</p>
<p>At the same time, the rapid progress in deep learning algorithms and their applications has accelerated the demand for high-performance computing platforms. In 1975, Moore projected a doubling of chip complexity every two years [<xref ref-type="bibr" rid="B3">3</xref>], but this trend has since approached a saturation in the possible density for conventional CMOS circuits. Recently, Graphics Processing Units (GPUs) have become the industry standard in scaling computing to meet demands, relying primarily on maximizing the use of parallel processing. AlexNet [<xref ref-type="bibr" rid="B4">4</xref>], introduced in 2012, was the first popular convolutional neural network architecture specifically developed for use on general-purpose graphics processing unit (GPGPU) platforms, following earlier proposed implementations such as [<xref ref-type="bibr" rid="B5">5</xref>, <xref ref-type="bibr" rid="B6">6</xref>]. Field-programmable gate array (FPGA) accelerators have also been introduced, including the Caffeine platform [<xref ref-type="bibr" rid="B7">7</xref>] in 2016, and the Microsoft Project Catapult [<xref ref-type="bibr" rid="B8">8</xref>] in 2017. NVIDIA has come to dominate the industry with their A100 accelerator introduced in 2020 [<xref ref-type="bibr" rid="B9">9</xref>]. Current benchmarks have shown that a single A100 can train a simple convolutional network to classify images from the CIFAR-10 dataset with 94% accuracy in just 3.29&#xa0;s [<xref ref-type="bibr" rid="B10">10</xref>].</p>
<p>But while these devices have shown great performance in terms of speed and scale, their energy demands are extreme: in 2023 alone, NVIDIA shipped 100,000 units, which will consume an average of 7.3&#xa0;TWh of electricity annually [<xref ref-type="bibr" rid="B11">11</xref>]. Currently, the majority of computational power demands come from data centers and are exacerbated by the increase in popularity of applications like artificial intelligence. Energy demands stem from the electricity supplying power (40%) and cooling requirements (40%), with the remainder attributed to associated compute infrastructure equipment. As a consequence, global electricity consumption by high-performance computing is expected to rise to a total range between 620 and 1,050&#xa0;TWh by 2026 [<xref ref-type="bibr" rid="B11">11</xref>]. This corresponds to an increase between 160-590&#xa0;TWh: roughly the annual demand of Sweden on the low end, or Germany on the high estimate.</p>
<p>Such trends reflect the fundamental limitations of acceleration through increased chip density and parallelism. As a result, in the search for ways of increasing scale to meet such application demands, research has begun to explore <italic>photonic accelerators</italic> as novel compute engines [<xref ref-type="bibr" rid="B12">12</xref>]. Photonic accelerators, also known as optical accelerators, are built on prior photonic technologies such as modulators, photo-detectors, and optical filters [<xref ref-type="bibr" rid="B13">13</xref>] which have been adapted to implement computing operations. This growth in interest is illustrated in <xref ref-type="fig" rid="F1">Figure 1</xref>, with a line plotting publications per year on photonic deep learning accelerators. Unlike traditional electronic components such as transistors and electronic switches, photonic accelerators utilize photons to process information. Photonic devices can make use of the properties of light to enable parallel processing and fast information transfer, with reduced energy consumption and greater efficiency per area.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>A timeline of milestones in accelerator development. The plot indicates the year in which each technology was introduced, showing the advancements in the 21st century which are foundational for practical photonic accelerators. The trajectory line indicates the number of publications per year on photonic deep learning accelerators in particular. Finally, the two horizontal lines indicate the projected throughput of conventional vs photonic accelerators, measured in tera-ops (TOPs) per area (mm<sup>2</sup>).</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g001.tif"/>
</fig>
<sec id="s1-1">
<title>1.1 Computing with light</title>
<p>The development of photonic accelerators has been driven by decades of innovations at the device and chip level of optical systems. These accelerators build upon foundational photonic technologies such as lasers, modulators, photodetectors, and optical filters. Many key developments in optical devices and integrated SiPh circuits have been introduced since the early 1980&#xa0;s, such as wavelength division multiplexing (WDM) filters [<xref ref-type="bibr" rid="B14">14</xref>&#x2013;<xref ref-type="bibr" rid="B16">16</xref>], Mach-Zehnder interferometer (MZI) modulators [<xref ref-type="bibr" rid="B17">17</xref>, <xref ref-type="bibr" rid="B18">18</xref>] and in-phase/quadrature (I/Q) modulators [<xref ref-type="bibr" rid="B19">19</xref>]. This evolution continued with the advent of smaller-sized Microring Resonators (MRRs), crucial in many optical filter designs, and high-speed or large bandwidth non-return-to-zero (NRZ) modulators [<xref ref-type="bibr" rid="B20">20</xref>]. Additionally, Pulse Amplitude Modulation with Four Levels (PAM4) modulation schemes have been explored, using ring resonators to increase the throughput per area of the device [<xref ref-type="bibr" rid="B21">21</xref>]. These ring resonators, possessing high-Q factors, have been engineered to function as switches, integrators, differentiators, and memory elements at both optical and terahertz (THz) frequencies.</p>
<p>The earliest optical accelerators could be traced in the assemblage of typical lab bench-top discrete optical components interconnected with long fiber spools intended to perform canonical mathematical functions [<xref ref-type="bibr" rid="B22">22</xref>, <xref ref-type="bibr" rid="B23">23</xref>]. One such important task is computing unitary operations, first demonstrated optically by Reck et al. [<xref ref-type="bibr" rid="B24">24</xref>] in 1994 using optical beam splitters, Fourier lenses, and light-emitting diode (LED) sources. This development laid the groundwork for subsequent advancements in integrated photonic computations using MZIs. Miller et al. [<xref ref-type="bibr" rid="B25">25</xref>&#x2013;<xref ref-type="bibr" rid="B27">27</xref>] showed that such MZI meshes could be self-configured to define a desired function, paving the way for building adaptive systems. Clements et al. [<xref ref-type="bibr" rid="B28">28</xref>] improved on the design with an alternative rectangular topology that achieves an equivalent computation using only half the optical depth. These landmark developments are plotted in the timeline of <xref ref-type="fig" rid="F1">Figure 1</xref>.</p>
<p>Optical computing has previously been viewed skeptically in applications that require large data storage and efficient flow control. However, current research demonstrates the capabilities of photonic accelerators on applications that are well-suited to the inherent advantages of optics. These applications include tasks with high parallelism, which can be efficiently computed by non-coherent optics through WDM, polarization diversity, and mode multiplexing [<xref ref-type="bibr" rid="B29">29</xref>]. Coherent approaches such as MZI circuits are more challenging to scale, raising concerns about high latency and insertion loss due to the longer physical length of the circuit [<xref ref-type="bibr" rid="B30">30</xref>], but MRRs present an alternative with better scalability and compactness. When light goes through ring resonators such as in 2 &#xd7; 2 switches, the drop port of the switch induces a time delay determined by the Q factor of the ring [<xref ref-type="bibr" rid="B31">31</xref>&#x2013;<xref ref-type="bibr" rid="B34">34</xref>]. This induced differential can be used in various ways to transmit information for computations. The latency can be tuned by inserting phase change materials (PCMs) as cladding, or cascading additional switches in tandem. The phase transition of the PCMs leads to appreciable alterations in their optical properties, controllable either electrically or optically [<xref ref-type="bibr" rid="B35">35</xref>, <xref ref-type="bibr" rid="B36">36</xref>]. This characteristic offers a notable advantage in power efficiency for programmable photonic devices, compared to electro-optic or thermo-optic methods [<xref ref-type="bibr" rid="B37">37</xref>, <xref ref-type="bibr" rid="B38">38</xref>].</p>
<p>Moreover, incorporating non-volatile PCMs as photonic devices enables optical memory storage and in-memory computing, achieved by transmitting optical input through the programmed device. For instance, optical memory in ring resonators has been studied using the Volterra series in microwave photonics [<xref ref-type="bibr" rid="B39">39</xref>]. The memory effect is modeled as a multidimensional impulse response in the time domain or Volterra kernels in the frequency domain. By using the ring resonator as a differentiator, it is possible to induce nonlinear mixing of multiple wavelengths to realize a frequency-dependent memory function.</p>
<p>More recently, these devices have been integrated to create energy-efficient, compact, and high-throughput computational accelerators. A comparative analysis of the theoretical maximum tera-operations per second per square millimeter (TOPs/mm<sup>2</sup>) for both electronic and photonic accelerators shows a clear advantage in the photonic domain.</p>
<p>To calculate the theoretical maximum TOPs/mm<sup>2</sup> for electronic accelerators, we consider the operational frequency (F), transistor density (D), and operations per cycle per transistor (O). The formula used is:<disp-formula id="equ1">
<mml:math id="m1">
<mml:msup>
<mml:mrow>
<mml:mtext>TOPs/mm</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msup>
<mml:mo>&#x003D;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>D</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>O</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>1</mml:mn>
<mml:msup>
<mml:mrow>
<mml:mn>0</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>12</mml:mn>
</mml:mrow>
</mml:msup>
<mml:mn>.</mml:mn>
</mml:math>
</disp-formula>
</p>
<p>For NVIDIA A100 [<xref ref-type="bibr" rid="B43">43</xref>], based on the TSMC 7&#xa0;nm node [<xref ref-type="bibr" rid="B44">44</xref>], the parameters are approximately: <italic>F</italic> &#x003D; 2&#xa0;GHz &#x003D; 2 &#xd7; 10<sup>9</sup>&#xa0;Hz, <italic>D</italic> &#x003D; 10<sup>8</sup> transistors/mm<sup>2</sup>, and <italic>O</italic> &#x003D; 2, which gives an estimate of 400 TOPs/mm<sup>2</sup>. However, due to constraints in practice, in the literature many electronic devices report a maximum efficiency of approximately 100 TOPs/mm<sup>2</sup> [<xref ref-type="bibr" rid="B46">46</xref>, <xref ref-type="bibr" rid="B47">47</xref>].</p>
<p>In contrast, for photonic accelerators, key parameters include the parallelism factor (P), component integration density (C), and efficiency factor (E). Their relationship is:<disp-formula id="equ3">
<mml:math id="m3">
<mml:msup>
<mml:mrow>
<mml:mtext>TOPs/mm</mml:mtext>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msup>
<mml:mo>&#x003D;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>C</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>E</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>1</mml:mn>
<mml:msup>
<mml:mrow>
<mml:mn>0</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>12</mml:mn>
</mml:mrow>
</mml:msup>
<mml:mn>.</mml:mn>
</mml:math>
</disp-formula>
</p>
<p>Taking the accelerator of Liu et al. as a conservative benchmark [<xref ref-type="bibr" rid="B45">45</xref>], representative current parameters are <italic>p</italic> &#x003D; 16384, <italic>C</italic> &#x003D; 10<sup>4</sup> components/mm<sup>2</sup>, and <italic>E</italic> &#x003D; 1, giving an estimate of 32 TOPs/mm<sup>2</sup>.</p>
<p>But while physical limitations increasingly constrain further enhancements in transistor density and operations per cycle for electronic accelerators advances in photonic technology may enable <italic>p</italic> &#x003D; 50,000, <italic>C</italic> &#x003D; 10<sup>5</sup> components/mm<sup>2</sup>, and <italic>E</italic> &#x003D; 1, potentially leading to 5000 TOPs/mm<sup>2</sup>, or 50 POPs/mm<sup>2</sup>&#x2014;performance measurable in peta operations per second.</p>
<p>Developments in photonic accelerators, alongside those in conventional hardware accelerators, are depicted in <xref ref-type="fig" rid="F1">Figure 1</xref>, which contrasts these comparative projected throughput capabilities of photonic computing versus electronic computing in terms of TOPs (tera-operations) per second normalized by processor area.</p>
<p>The significantly higher level of projected TOPS/mm<sup>2</sup> for photonic systems is attributed to the efficient parallelism achieved through utilizing multiple wavelengths, coupled with a smaller footprint per wavelength.</p>
<p>In silicon nitride (SiN) photonics-based devices, the area of one MAC unit cell is 285 &#xd7; 354&#xa0;<italic>&#x3bc;m</italic>
<sup>2</sup> [<xref ref-type="bibr" rid="B48">48</xref>, <xref ref-type="bibr" rid="B49">49</xref>]. This, when operating at 12&#xa0;GHz with 4 input vectors via WDM, corresponds to a compute density of 1.2&#xa0;TOPS/mm<sup>2</sup>. If silicon-on-insulator (SOI) MRR devices are used instead with a nominal bend radius of 5<italic>&#x3bc;m</italic>, the area of the MAC unit cell could be reduced to less than 30 &#xd7; 30&#xa0;<italic>&#x3bc;m</italic>
<sup>2</sup>, increasing the compute density to 420 TOPS/mm<sup>2</sup> per input channel [<xref ref-type="bibr" rid="B50">50</xref>, <xref ref-type="bibr" rid="B51">51</xref>]. In-memory-computing photonic tensor cores show predicted compute density and compute efficiencies of 880 TOPS/mm<sup>2</sup> and 5.1 TOPS/W for a 64 &#xd7; 64 crossbar core at 25&#xa0;GHz clock speed [<xref ref-type="bibr" rid="B52">52</xref>]. Compared with digital electronic accelerators (ASIC and GPU), the photonic core has 1 to 3 orders of magnitude improvement in both compute density and efficiency. Overall, this comparison underscores the advancements and potential of photonic technologies in achieving higher throughput and efficiency in computing. This makes it a competitive candidate for application in the context of neural network processing and deep learning acceleration.</p>
</sec>
<sec id="s1-2">
<title>1.2 Photonics for deep learning</title>
<p>Researchers have been interested in optical implementations of neural networks since the 1980s [<xref ref-type="bibr" rid="B40">40</xref>], for instance, exploring image recognition by the use of nonlinear joint transform correlators [<xref ref-type="bibr" rid="B22">22</xref>], and implementing Hopfield neural networks [<xref ref-type="bibr" rid="B41">41</xref>, <xref ref-type="bibr" rid="B42">42</xref>]. Since then, many innovations have stemmed from advancements in photonic tensor cores, in-memory computing, and hybrid co-processors [<xref ref-type="bibr" rid="B35">35</xref>, <xref ref-type="bibr" rid="B53">53</xref>&#x2013;<xref ref-type="bibr" rid="B57">57</xref>]. For instance, in deep learning inference, trained weights may not require frequent updates or any at all, making non-volatile analog memory advantageous. This can be achieved using PCMs, either optically [<xref ref-type="bibr" rid="B58">58</xref>, <xref ref-type="bibr" rid="B59">59</xref>] or electronically [<xref ref-type="bibr" rid="B60">60</xref>, <xref ref-type="bibr" rid="B61">61</xref>]. On the other hand, a real-time neural network can be established by using digital electronic drivers with photonic-compatible firmware. Neuron behavior can be replicated through a hybrid of well-modeled electronic nonlinearities and optical systems that have negligibly low losses. In those systems, the active components consist of photodetectors (PDs) and modulators that inject or deplete carriers in response to an induced electric field [<xref ref-type="bibr" rid="B62">62</xref>, <xref ref-type="bibr" rid="B63">63</xref>].</p>
<p>Photonic computing and its use in artificial intelligence applications can be viewed from a multitude of perspectives, many of which have been previously explored in reviews. Various reviews have been devoted to photonic analog computing broadly, such as Stroev and Berloff [<xref ref-type="bibr" rid="B64">64</xref>]. Huang et al. [<xref ref-type="bibr" rid="B65">65</xref>] provide a survey of design factors in neuromorphic computing, and discuss the role of photonic processing for implementing aspects such as interconnects, linear vs nonlinear operations, and memory, as well as presenting use cases in communications, nonlinear programming, and cryptography. Wu et al. [<xref ref-type="bibr" rid="B66">66</xref>] review analog optical computing based on integrated photonics, diffractive networks, and hybrid optoelectronic designs applied specifically to three classes of machine learning models: feed-forward networks, spiking neural networks, and reservoir computing.</p>
<p>In this review, we present a concise overview of the photonic accelerator landscape to provide context for photonic deep learning accelerators (PDLAs), and provide some background on elements of the compute operations in deep neural network architectures that are mapped onto photonic implementations. We focus on Silicon Photonics Integrated Circuit (Si PIC) accelerators, as this modality can be considered more practical for near-term use given its level of technical advancement, cost-effectiveness, and compatibility with conventional CMOS hardware. Our analysis seeks to unify low-level design considerations in implementing PDLAs with a broader perspective on application. The paper is organized as follows: in <xref ref-type="sec" rid="s2">Section 2</xref>, we give context on the broader area of photonic accelerator design: physical <italic>modalities</italic>, as well as analog and digital <italic>compute paradigms</italic>. <xref ref-type="sec" rid="s3">Section 3</xref> provides an overview of the computational building blocks in deep learning, and indicates the roles that photonic accelerators can play in neural network models. In <xref ref-type="sec" rid="s4">Section 4</xref>, we highlight specific approaches to PDLA design with representative examples from the literature. Finally, <xref ref-type="sec" rid="s5">Section 5</xref> indicates ongoing challenges in implementing PIC-based systems and promising further directions for research, with key takeaways for both photonics and deep learning practitioners.</p>
</sec>
</sec>
<sec id="s2">
<title>2 Photonic accelerators</title>
<p>Photonic principles can be used for accelerated computing in many ways, so we first provide context on the primary physical <italic>modalities</italic> used in a photonics processor. Those devices can also operate in both analog and digital <italic>computing paradigms</italic>, and we provide examples of each approach. <xref ref-type="fig" rid="F2">Figure 2</xref> shows this schema of physical and computational properties of photonic accelerators.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Schema of structural and computational factors in photonic accelerator design.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g002.tif"/>
</fig>
<sec id="s2-1">
<title>2.1 Physical modalities</title>
<p>
<italic>Optical Processing Units (OPUs)</italic> are photonic devices used for computing tasks, efficiently performing a broad range of mathematical and logical tasks crucial for applications such as deep learning. These devices leverage optics instead of electronics, in contrast with traditional CMOS processors such as CPUs, GPUs, and TPUs. OPUs have demonstrated scalability in facilitating acceleration within standard computing frameworks [<xref ref-type="bibr" rid="B67">67</xref>]. High-bandwidth optical interconnects are central to optical data transmission accelerators, and recent advances here have focused on increasing data rates, decreasing power consumption, and achieving higher reliability [<xref ref-type="bibr" rid="B52">52</xref>, <xref ref-type="bibr" rid="B68">68</xref>]. OPUs can be based on three main modalities: integrated optics, quantum optics, and free space optics.</p>
<sec id="s2-1-1">
<title>2.1.1 Integrated circuit OPUs</title>
<p>Photonics Integrated Circuits (PICs), the predominant form of OPUs, are engineered for efficiency in operations such as matrix multiplication and convolution [<xref ref-type="bibr" rid="B69">69</xref>]. Integrated optical processors have been demonstrated for implementing matrix-vector multiplications at Gb/s processing rates [<xref ref-type="bibr" rid="B70">70</xref>&#x2013;<xref ref-type="bibr" rid="B72">72</xref>]. Companies like Lightmatter<xref ref-type="fn" rid="fn3">
<sup>1</sup>
</xref>, Lightelligence<xref ref-type="fn" rid="fn4">
<sup>2</sup>
</xref>, Luminous<xref ref-type="fn" rid="fn5">
<sup>3</sup>
</xref> are developing photonics ICs for low-power multiply-and-accumulate (MAC) computations which significantly outperform conventional digital and analog electronics.</p>
<p>Adaptive and reconfigurable OPUs also represent an emerging subgroup with the ability to dynamically alter processing parameters, an essential requirement for many machine learning use cases [<xref ref-type="bibr" rid="B75">75</xref>]. Programmable OPUs eliminate the need for physical hardware modifications, ensuring cost-effectiveness and resource efficiency. Harris et al. [<xref ref-type="bibr" rid="B76">76</xref>] reviewed progress made in Programmable Nanophotonic Processors (PNPs), which employ both classical and quantum information processing. Bogaerts et al. [<xref ref-type="bibr" rid="B77">77</xref>] present a survey of the photonic building blocks, as well as discussing the necessary control structures and application-level considerations, for instance highlighting the need for developing descriptive languages similarly to FPGA programming.</p>
<p>An important approach in reprogrammable device design is the use of phase-change materials (PCMs). For example, Wu et al. [<xref ref-type="bibr" rid="B35">35</xref>] propose a compact, programmable waveguide mode converter based on a Ge<sub>2</sub>Sb<sub>2</sub>Te<sub>5</sub> (GST-enhanced) phase-gradient metasurface. The converter uses changes in the refractive index of GST to control the waveguide spatial modes up to 64 levels. This contrast represents the matrix elements, with a 6-bit resolution to perform matrix-vector multiplication in convolutional neural networks. The design featured high programming resolution and was used to construct a photonic kernel using an array of such phase-change metasurface mode converter (PMMC) devices, enabling an optical convolutional neural network to be designed for image processing and recognition tasks. The authors use nanogap-enhanced potential for a wide range of optical functions, making them suitable for large-scale optical computing and neuromorphic photonics.</p>
<p>Innovations in this category also address the issue of noise through advanced noise reduction and error correction techniques, which are important properties in supporting the accuracy and reliability of machine learning computations [<xref ref-type="bibr" rid="B78">78</xref>]. The researchers in [<xref ref-type="bibr" rid="B79">79</xref>&#x2013;<xref ref-type="bibr" rid="B82">82</xref>] offer a comprehensive review of PCMs in non-volatile photonic applications. They highlight the retention of the optical state of a material without the need for continuous power supply, and the potential for low-energy operation due to the efficient transformation between amorphous and crystalline states, providing a pathway to highly reconfigurable photonic devices.</p>
</sec>
<sec id="s2-1-2">
<title>2.1.2 Quantum OPUs</title>
<p>Quantum OPUs represent another approach to OPU design. These devices have been previously developed and applied in the context of communications [<xref ref-type="bibr" rid="B83">83</xref>, <xref ref-type="bibr" rid="B84">84</xref>]. Quantum OPUs can implement compute tasks on very small scales. For example, quantum dots are devices that have small dimensions of a few nanometers. Quantum Dot (QD)&#x2013;based OPUs incorporate quantum dots, nanoscale semiconductor particles with dimensions of several nanometers, to enhance OPU functionality. Semiconductor QDs represent a type of zero-dimensional, quantum-confined device which exhibits distinct electronic and optical characteristics. The three-dimensional quantum confinement within QDs leads to the total localization of carriers, producing a discrete spectrum characterized by a <italic>&#x3b4;</italic>-function-like density of states [<xref ref-type="bibr" rid="B85">85</xref>]. The precision control afforded by these quantum dots over photon emission and absorption translates to more effective processing tailored for specific machine learning tasks, thereby expanding the versatility of photonic processing applications [<xref ref-type="bibr" rid="B86">86</xref>].</p>
<p>Lingnau et al. [<xref ref-type="bibr" rid="B87">87</xref>] furthered the domain with the use of coupled quantum well devices on-chip, highlighting their potential in creating excitable neuromorphic networks [<xref ref-type="bibr" rid="B88">88</xref>]. Present a PIC consisting of quasi-single-mode slotted Fabry&#x2013;P&#xe9;rot lasers coupled via an actively pumped waveguide. This research shows how quantum optics can enable a variety of controllable excitable states, including dual-state excitability and dual-state bursting mixed-mode oscillations. A state-of-the-art large-scale integrated quantum photonic circuit [<xref ref-type="bibr" rid="B89">89</xref>] has been successfully demonstrated in silicon, boasting 16 waveguide spirals, 93 reconfigurable thermo-optical phase shifters, 122 MMIs, 64 grating couplers, and 376 crossings. This reconfigurable device showcased its capabilities in generating, manipulating, and managing (GMM) entangled states directly on the chip.</p>
<p>Quantum photonics can allow for implementing quantum algorithms on an integrated device, for instance implementing Shor&#x2019;s algorithm to factorize 15 into 3 and 5 [<xref ref-type="bibr" rid="B90">90</xref>]. This system comprises a Quantum Fourier Transform subsystem and a two-qubit controlled NOT gate. Variants of quantum photonic algorithms akin to these have been employed in solving a standard eigenvalue problem [<xref ref-type="bibr" rid="B91">91</xref>], as well as in the implementation of graph-theoretic algorithms utilizing a SiPh quantum walk processor [<xref ref-type="bibr" rid="B92">92</xref>]. However, realizing these quantum-enhanced accelerators presents many technical challenges and feasibility questions [<xref ref-type="bibr" rid="B93">93</xref>]. Processing single photons in large quantities requires high-speed, low-loss optical switches like lithium niobate and barium titanate. Achieving the complete integration of quantum circuits, including sources and detectors, remains an unresolved endeavor.</p>
</sec>
<sec id="s2-1-3">
<title>2.1.3 Free-space photonics</title>
<p>Free-space optics represents a pivotal modality in optical computing, diverging from traditional silicon-based mediums to leverage plane light propagation in free space. This approach, as Hsu et al. [<xref ref-type="bibr" rid="B94">94</xref>] highlights, exploits additional degrees of freedom such as polarization, diffraction, and orbital angular momentum (OAM), making it particularly suited to tasks involving imaging data and computer vision applications. The use of diffraction for manipulating incident light, as demonstrated by Zhu et al. [<xref ref-type="bibr" rid="B95">95</xref>], and the implementation of a Laguerre-Gaussian mode sorter (LGms) for super-multimode (de)multiplexing in optical communications by Fontaine et al. [<xref ref-type="bibr" rid="B96">96</xref>], underscore the versatility and potential of free-space optics in enhancing optical computing capabilities.</p>
<p>Deep diffractive neural networks (D<sup>2</sup>NNs) stand as a notable application of free-space optics. Lin et al.&#x2019;s D<sup>2</sup>NN uses passive diffractive layers to implement transforms, though it lacks rapid programmability [<xref ref-type="bibr" rid="B97">97</xref>]. Another D<sup>2</sup>NN design employed orbital angular momentum (OAM) to adjust the phase and amplitude across multiple diffractive screens, enabling the manipulation of light beams&#x2019; wavefronts for a trainable network architecture. Hamerly et al. advanced the application of free-space optics in optical computing by employing quantum photoelectric multiplication to implement matrix-vector products through coherent detection [<xref ref-type="bibr" rid="B98">98</xref>]. This method not only allows the optical encoding of weights and inputs but also supports the reprogramming and training of the accelerator. Capable of operating at GHz speeds with sub-attojoule energy per MAC, this accelerator scales to larger networks with <italic>N</italic> &#x2265; 106 neurons. Another demonstration of D<sup>2</sup>NNs is reported in [<xref ref-type="bibr" rid="B99">99</xref>, <xref ref-type="bibr" rid="B100">100</xref>] with programmable optoelectronic devices as well as additional variants such as D-NIN-1, and D-RNN. Such capabilities indicate the increasing potential of free-space devices in realizing practical, large-scale applications in areas like deep learning, marking a departure from fully integrated photonic processors.</p>
<p>Free-space devices show promise for large scalability, as shown by the LightOn OPU [<xref ref-type="bibr" rid="B73">73</xref>] which can operate at 50 TOPS/watt with input vector dimensions of 1 million &#xd7; 2 million. This OPU can accelerate randomized numerical linear algebra algorithms by implementing very large random matrices optically. It shows how optical properties such as scattering can circumvent the limitations of a von Neumann architecture by performing high-dimensional operations in a single computational step, reducing the effective complexity from <italic>O</italic>(<italic>n</italic>
<sup>2</sup>) to <italic>O</italic>(1). Moreover, the exploration of complex analog computations in free space, as investigated by Cordaro et al., further exemplifies the innovative uses of this technology [<xref ref-type="bibr" rid="B101">101</xref>]. Their work on using a silicon metasurface-based platform to solve Fredholm integral equations of the second kind illustrates the broad applicability and the advanced computational possibilities enabled by free-space optics. Collectively, these developments not only underscore the technological advancements in free-space optical computing, but also highlight its expanding role in addressing sophisticated computational challenges.</p>
</sec>
</sec>
<sec id="s2-2">
<title>2.2 Computing paradigms</title>
<p>Analog processors leverage the continuous-time and space properties of light to perform computations, whereas digital photonic accelerators use digital encoding. This flexibility offers two approaches to processing photonic signals and to designing photonic accelerators for machine learning tasks.</p>
<sec id="s2-2-1">
<title>2.2.1 Analog optical processing</title>
<p>Analog Optical Processing Units (A-OPUs) [<xref ref-type="bibr" rid="B64">64</xref>] use the continuous values generated by the physical functionality of the device by reading them out as computation results, to perform operations like weighted summation in an energy-efficient manner. This is particularly useful in scientific simulations and optimization problems, where continuous solutions are desired. <xref ref-type="fig" rid="F3">Figure 3</xref> shows an example of an A-OPU suited to solve partial differential equations (PDEs) and ordinary differential equations (ODEs) [<xref ref-type="bibr" rid="B102">102</xref>]. When the temporal frequency of the input signal is near the resonant frequency of the phase-shifted Distributed Feedback Semiconductor Optical Amplifier (DFB-SOA), the resultant transfer function becomes equal to that of a first-order linear ODE. Adjusting the injection current at the input tunes the constant coefficient of this ODE. In this way the phase-shifted DFB-SOA can be used to implement a photonic ODE solver by controlling the injection current.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Schematic of an analog photonic ODE solver. <bold>(A)</bold> When injection current <inline-formula id="inf3">
<mml:math id="m10">
<mml:mrow>
<mml:mi mathvariant="italic">I</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is below the lasing threshold, the constant coefficient of the ODE can be tuned by modulating the current. <bold>(B)</bold> The DFB-SOA can be biased to operate at the lasing mode and connected to an optical filter, achieving temporal intensity differentiation. Reproduced without changes under terms of the CC-BY license from [<xref ref-type="bibr" rid="B102">102</xref>], Li et al. 2016, &#xa9; Springer Nature.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g003.tif"/>
</fig>
<p>Analog photonic processing has also been applied to reservoir computing (RC). Originating from concepts in liquid-state machines and echo-state networks, RC is a type of machine learning framework which maps inputs into a fixed non-linear system, known as a &#x201c;reservoir,&#x201d; then processes this information through a trainable readout mechanism to produce the model output [<xref ref-type="bibr" rid="B103">103</xref>]. The reservoir can be implemented in many ways, and A-OPU devices are increasingly explored as analog reservoirs, showing success when applied to time-series data processing and pattern recognition tasks [<xref ref-type="bibr" rid="B104">104</xref>&#x2013;<xref ref-type="bibr" rid="B107">107</xref>]. Further, A-OPUs have also played a role in quantum photonic processing, as seen in Continuous-variable Quantum Optical Processors (CQOPs) [<xref ref-type="bibr" rid="B84">84</xref>, <xref ref-type="bibr" rid="B108">108</xref>&#x2013;<xref ref-type="bibr" rid="B111">111</xref>]. The analog approach uses the inherent properties of photon behavior to efficiently perform quantum simulations or produce solutions to optimization problems [<xref ref-type="bibr" rid="B64">64</xref>].</p>
</sec>
<sec id="s2-2-2">
<title>2.2.2 Digital optical processing</title>
<p>Digital Optical Processing Units (D-OPUs), on the other hand, use discrete photonic signals for computation and processing [<xref ref-type="bibr" rid="B112">112</xref>, <xref ref-type="bibr" rid="B113">113</xref>]. D-OPUs are often designed around enabling typical computing operations like binary logic and bit manipulation, but in a fast and efficient manner using the optical domain. Gostimirovic et al. [<xref ref-type="bibr" rid="B114">114</xref>] proposed a hybrid photonic-electronic circuitry for a digital logic architecture using ultra-compact vertical pn junctions based on microdisk switches. With higher &#x394;<italic>&#x3bb;</italic>/<italic>V</italic>, where <italic>V</italic> is the voltage, they used wavelength-division multiplexing to implement NAND, NOR, and XNOR operations with a single MRR switch. The gates are then expanded to explore complex CMOS-compatible blocks such as adders, encoders, and decoders. Several aspects of optical logic computing have also been explored using semiconductor optical amplifiers (SOAs) [<xref ref-type="bibr" rid="B115">115</xref>, <xref ref-type="bibr" rid="B116">116</xref>]. Many mathematical operations can be implemented using Binary Photonic Arithmetic (BPA) where photonic accelerators perform binary arithmetic operations using discrete optical signals [<xref ref-type="bibr" rid="B117">117</xref>]. Digital photonic data transmission has also emerged in optical interconnects for data compression, multiplexing, and encoding. These technologies facilitate digital data handling between processing units and memory components in high-performance computing clusters.</p>
<p>In addition to standard bit operations, quantum photonic devices can be used to achieve qubit behavior to facilitate quantum algorithms. Such Quantum Digital Optical Processors (QDOP) [<xref ref-type="bibr" rid="B118">118</xref>] can reach ultrafast (1&#xa0;Tb/s) speeds for optical logic operations [<xref ref-type="bibr" rid="B119">119</xref>]. In this context, quantum dot (QD) SOAs have advantages such as minimal crosstalk between adjacent wavelength channels due to QD isolation, which suppressed carrier transfer between dots, and utilization of the cross gain modulation (XGM) effect between two wavelength channels [<xref ref-type="bibr" rid="B120">120</xref>, <xref ref-type="bibr" rid="B121">121</xref>]. These QDOP units would enable quantum computations and algorithms that work with digital quantum information, facilitating quantum-enhanced machine learning algorithms.</p>
</sec>
</sec>
</sec>
<sec id="s3">
<title>3 Photonic deep learning fundamentals</title>
<p>Photonic accelerators for deep learning are built on the functionalities of photonic devices highlighted in <xref ref-type="sec" rid="s2">Section 2</xref>. The fundamental goal is to perform the intensive computations required by deep neural networks efficiently and at high speed. Neural networks are built out of linear products and nonlinear special functions. Deep networks include many layers of these operations, which results in their computational expense. Accelerator design can target different components of a network, from the lowest level of mathematical operations to higher-level architecture blocks. Here, we present an overview of the main neural network components and the ways that they are translated to photonic implementations, along with some ways in which performance considerations must be reinterpreted.</p>
<sec id="s3-1">
<title>3.1 MAC operations in neural networks</title>
<p>The bulk of a network&#x2019;s computation comes from the matrix multiplications present in layer transforms, and one way to assess network complexity is to count the number of multiply-accumulate (MAC) operations required to evaluate the full network on a given input. For a modified state <italic>a</italic>&#x2032; and a given accumulation variable <italic>a</italic>, a MAC operation can be written as <inline-formula>
<mml:math id="i3">
<mml:msup>
<mml:mi>a</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
<mml:mo>&#x2190;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x002B;</mml:mo>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>w</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>.</mml:mo>
</mml:math>
</inline-formula>
</p>
<p>In the general case of a linear layer in a network, the action of the layer on an input consists of a weighted sum<disp-formula id="e2">
<mml:math id="m7">
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
<mml:mo>&#x003D;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mi>i</mml:mi>
</mml:munder>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:msub>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x002B;</mml:mo>
<mml:msub>
<mml:mi>b</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:mo>.</mml:mo>
</mml:math>
<label>(1)</label>
</disp-formula>
</p>
<p>&#x201c;Neurons&#x201d; <italic>x<sub>i</sub>
</italic> from layer <italic>i</italic> transfer signals to neuron <italic>x<sub>j</sub>
</italic> in the following layer <italic>j</italic> through connection weights <italic>w</italic>
<sub>
<italic>ij</italic>
</sub>, linking a set of input and output variables. <italic>b</italic>
<sub>
<italic>j</italic>
</sub> is a &#x201c;bias&#x201d; offset for translation, making it an affine transform. <italic>f</italic>{&#x22c5;} represents a discriminatory nonlinear &#x201c;activation&#x201d; function [<xref ref-type="bibr" rid="B122">122</xref>, <xref ref-type="bibr" rid="B123">123</xref>]. In a typical network, this is chosen to be either a sigmoid-shaped function, such as the logistic or hyperbolic tangent functions, or a ramp-shaped function, such as the rectified linear unit (<italic>ReLU</italic> &#x003D; max{0, <italic>x</italic>}). The output variables <italic>x</italic>
<sub>
<italic>j</italic>
</sub> are often referred to as the &#x201c;activations.&#x201d; The weighted sum of Eq. <xref ref-type="disp-formula" rid="e2">1</xref> forms a set of parallel MAC operations and is thus computed as a matrix multiplication of size <italic>N</italic> &#xd7; <italic>M</italic> to convert an input of size <italic>M</italic> to an output of size <italic>N</italic>, and in terms of computational complexity often accounted for as <inline-formula id="inf1">
<mml:math id="m8">
<mml:mi>O</mml:mi>
<mml:mfenced close=")" open="(">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:math>
</inline-formula>, given that in practice the input and output size of internal layers are typically of similar magnitude.</p>
<p>Convolutional neural networks, on the other hand, act on windows of the input tensor, making use of the locality of information in data. As a result, they are especially suitable for tasks on images and other natural signals. Conceptually, a 2D convolution layer takes in a 3D input tensor of size (<italic>H</italic> &#xd7; <italic>W</italic> &#xd7; <italic>C</italic>
<sub>
<italic>in</italic>
</sub>) and a 4D kernel tensor of size (<italic>C</italic>
<sub>
<italic>in</italic>
</sub> &#xd7; <italic>C</italic>
<sub>
<italic>out</italic>
</sub> &#xd7; <italic>k</italic>
<sub>0</sub> &#xd7; <italic>k</italic>
<sub>1</sub>), and outputs a 3D tensor of size (<italic>H</italic> &#xd7; <italic>W</italic> &#xd7; <italic>C</italic>
<sub>
<italic>out</italic>
</sub>). Overall, the layer must apply the kernel transform to all <italic>k</italic>
<sub>0</sub> &#xd7; <italic>k</italic>
<sub>1</sub> windows of the input, multiplying them together and summing the values in a convolution operation. In practice, kernel windows are usually square and relatively small (width <inline-formula id="inf2">
<mml:math id="m9">
<mml:mo>&#x003c;</mml:mo>
</mml:math>
</inline-formula> 10). However, the input and output channel numbers may be in the hundreds (e.g. up to 512 in VGG [<xref ref-type="bibr" rid="B124">124</xref>]). In CNNs, the activation functions are often followed by a pooling operation over windows of the output, which may consist of further MACs (as in average pooling), or of another nonlinear function (as in maximum pooling).</p>
<p>Computationally, there are many ways of formulating this multiple-channel, multiple-kernel convolution as generalized matrix-matrix multiplication (GEMM) suitable for modern hardware [<xref ref-type="bibr" rid="B125">125</xref>]. The <monospace>im2col</monospace> algorithm is often used as a conceptual basis, vectorizing the input such that its values are duplicated for multiplication with the kernel. This naive construction results in a matrix multiplication between matrices of size <italic>M</italic> &#xd7; (<italic>ck</italic>
<sup>2</sup>) and (<italic>HW</italic>) &#xd7; (<italic>ck</italic>
<sup>2</sup>), as depicted in <xref ref-type="fig" rid="F4">Figure 4</xref> [<xref ref-type="bibr" rid="B126">126</xref>]. Here the desired number of output channels is reflected in the value <italic>M</italic>. Modern GPU implementations derive their efficiency from optimizations such as re-using intermediate results and reducing the amount of matrix reshaping. They also apply virtual memory strategies so that re-used values are never physically duplicated in memory.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>A depiction of the basic im2col formulation of multi-channel multi-kernel convolution as a generalized matrix-matrix multiplication (GEMM). Reproduced from [<xref ref-type="bibr" rid="B105">105</xref>], &#xa9; IEEE.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g004.tif"/>
</fig>
</sec>
<sec id="s3-2">
<title>3.2 Photonic network principles</title>
<p>Designing PIC accelerators for deep learning relies on translating photonic capabilities to these essential building blocks of neural networks. Accelerators can target the linear operations of feedforward and convolution layers through fast photonic multiply and accumulate methods. They can also target nonlinear functions through switching and modulating.</p>
<p>As an illustration, in their groundbreaking work, Shen et al. [<xref ref-type="bibr" rid="B127">127</xref>] laid out a construction of how photonic elements can be mapped onto the components of a feedforward neural network for an all-optical procedure. The linear operation of matrix multiplication can be formulated as a unitary operation and readily implemented in programmable photonic circuits (PPC), where phase shifters can tune the optical paths, allowing reconfiguration of neural network weights. Nonlinear activations can then be performed using optical switching elements such as saturable absorbers. They also observe how in-situ training of such an accelerator can be realized photonically not by the backpropagation algorithm standard to digital NNs, but rather by forward propagation and finite differencing to directly obtain the gradient of each parameter. Following the PPC approach of Shen et al., many accelerators choose to implement linear operations photonically using PPC for universal unitary operations. More recently, as an alternative to general MZI mesh designs, Shokraneh et al. [<xref ref-type="bibr" rid="B128">128</xref>] designed a &#x201c;diamond&#x201d; mesh structure specifically optimized for use in neural networks.</p>
<p>In contrast with coherent, PPC-based designs, the other main concept for linear operations in photonic accelerators is to leverage non-coherent photonics through wavelength division multiplexing (WDM) for parallel operations at scale. The broadcast-and-weight protocol of Tait et al. [<xref ref-type="bibr" rid="B129">129</xref>] applied the analogy of the broadcast-and-select WDM protocol by observing the similar network connectivity of neurons between layers. Tunable filter banks based on microring resonators (MRRs) can thus be used similarly to how wavelength demultiplexers are realized in conventional digital interconnects. While the protocol was originally introduced for linear network layers, conceptually this extends naturally to convolutional layers, as a linear layer is equivalent to a convolutional layer with a &#x201c;1 &#xd7; 1&#x201d; kernel. Feldmann et al. [<xref ref-type="bibr" rid="B130">130</xref>] have since demonstrated a photonic tensor core that combines the abilities of microcombs and phase-change materials to realize efficient encoding of data and kernels, respectively. Movement of data is minimized with in-memory photonic MAC operations and reduces footprint cost by multiplexing within a single core. Meanwhile, Xu et al. [<xref ref-type="bibr" rid="B131">131</xref>] introduced a convolutional accelerator that emphasizes maximized input size capacity, handling full-resolution images of 500 &#xd7; 500 pixels by making use of both time and wavelength interleaving.</p>
<p>Photonic devices are naturally suited to the linear nature of matrix multiplication, but it is also possible to implement all-optical activations with optical switching implemented for instance in the action of a saturable absorber or nanocavities, as suggested by [<xref ref-type="bibr" rid="B127">127</xref>]. Other possibilities include using carrier effect in MRR, or state changes in a material as in a structural phase transition [<xref ref-type="bibr" rid="B65">65</xref>]. In addition, some accelerators implement pooling operations photonically, for instance with ring modulators [<xref ref-type="bibr" rid="B132">132</xref>], or MMIs [<xref ref-type="bibr" rid="B133">133</xref>]. However, the power consumption required to trigger activation switches and to maintain a sufficient signal-to-noise ratio at receiving photodetectors can dominate otherwise passive multiplication steps [<xref ref-type="bibr" rid="B127">127</xref>]. As such, many accelerator designs compute these functions in a hybrid optoelectronic manner, converting the output to the electronic domain between multiplication layers.</p>
<p>In addition to handling the arithmetic intensity of deep learning applications, memory implementation is an essential consideration when developing practical hardware accelerators. One important implementation is memristors. Memristors, or resistance switches, were first proposed theoretically as the completion of the three other &#x201c;fundamental&#x201d; electrical components: resistors, capacitors, and inductors [<xref ref-type="bibr" rid="B134">134</xref>]. The internal state of a memristor is a function of the history of current and/or voltage which has passed through it [<xref ref-type="bibr" rid="B135">135</xref>]. Devices that contain &#x201c;crossbar&#x201d; arrays of connected memristors have been successfully applied in deep learning applications. A noteworthy example is the ISAAC accelerator [<xref ref-type="bibr" rid="B136">136</xref>], which introduced the use of electronic memristive crossbar arrays. Since then, optical memristors have shown improved efficiency over electronic versions in accelerators. Mao et al. [<xref ref-type="bibr" rid="B137">137</xref>] provide a comprehensive overview of how practical memristor behavior can be implemented with photonic elements, and highlights how memristors can have various functionalities for light detection, data storage, and in-memory computing. Choi et al. [<xref ref-type="bibr" rid="B138">138</xref>] demonstrate a model of in-memory processing that can be realized by photonics integrated circuits using coupled resonators, where the coupled memristive quantities are the intensity distribution and optical coherence. They indicate that their design is scalable to neural network applications.</p>
</sec>
<sec id="s3-3">
<title>3.3 Performance considerations</title>
<p>When translating neural network computation to alternative hardware, it can be challenging to make direct comparisons in different aspects of performance. In conventional hardware, the layer transform is considered the primary MAC hardware bottleneck as layer size grows [<xref ref-type="bibr" rid="B139">139</xref>]. <xref ref-type="fig" rid="F5">Figure 5</xref> indicates the requirements in hardware which performs MACs individually and does not compute in memory. In this case, network MACs can be counted uniformly, and for modern networks such as Vision Transformer or ResNet, this can reach 500 billion MACs in a single forward pass<xref ref-type="fn" rid="fn6">
<sup>4</sup>
</xref>. However, a full optical matrix multiplication can be performed, in principle, in a single step, without consuming any power, independent of the matrix size [<xref ref-type="bibr" rid="B127">127</xref>, <xref ref-type="bibr" rid="B140">140</xref>]. The main sources of energy consumption or latency are generally shifted to aspects of transmission, modulation, and detection, performed by various components in the device [<xref ref-type="bibr" rid="B139">139</xref>], so photonic device architectures must make tradeoffs in balancing these factors.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>An illustration of the signal pathway required for MAC in a modern chip. The passing of information occurs between MAC processors performing <italic>a</italic> &#x002B; (<italic>w</italic> &#xd7; <italic>x</italic>), memory caches, and non-linear operations <italic>f</italic>{&#x22c5;}. Reproduced with permission, from [<xref ref-type="bibr" rid="B139">139</xref>], Nahmias et al. 2020 &#xa9; IEEE.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g005.tif"/>
</fig>
<p>As a result, the complexity of photonic MAC operations must be conceptualized differently than in conventional hardware. In the photonic case, &#x201c;complexity&#x201d; is no longer tied to algorithmic complexity in terms of counting individual multiply-accumulate steps. As Miscuglio et al. state, &#x201c;one must distinguish between the complexities of the computational algorithm vs that of the system&#x2019;s <italic>execution time</italic>&#x201d; [<xref ref-type="bibr" rid="B140">140</xref>]. By comparison, it is important to note that GPUs are still bound by the <italic>O</italic>(<italic>N</italic>
<sup>2.8</sup>) (Strassen [<xref ref-type="bibr" rid="B141">141</xref>]) or <italic>O</italic>(<italic>N</italic>
<sup>2.373</sup>) (Coppersmith-Winograd [<xref ref-type="bibr" rid="B142">142</xref>]) complexity of matrix multiplication algorithms, and their optimization is in reduced system execution time due to parallelism, value re-use, and minimized I/O cost. In order to make comparisons, a more appropriate frame is to think of &#x201c;effective&#x201d; MACs per time. For instance, when we say that a photonic operation is &#x201c;O(1),&#x201d; we mean that the entire computation is executed in a single &#x201c;atomic&#x201d; computing operation. In a passive component, this can effectively be the speed of light propagating through the medium. This is also why compute <italic>density</italic> becomes a more important metric to consider, as photonic components may individually be larger, but a single component can implement many &#x201c;effective&#x201d; MACs. As a result, many accelerators report normalized performance statistics in terms of operations per area.</p>
<p>Another distinction arises particularly in analog photonic accelerators in the way that &#x201c;bit precision&#x201d; is translated to photonic hardware. As discussed by Shiflett et al. [<xref ref-type="bibr" rid="B143">143</xref>], &#x201c;While we use the terminology &#x2018;bits of precision&#x2019; for analog photonic computation, what we are actually describing is the log<sub>2</sub> of the number of <italic>separable optical power amplitudes at the output</italic>.&#x201d; Numerical precision becomes reliant on the signal-to-noise ratio of transmission among device components. This presents a source of energy overhead as for instance the power of input lasers must be increased in order to increase this ratio. In the case of MRR-based designs, a tradeoff between multiplexing parallelism and numerical precision may also arise, through the power cross-coupling coefficient <italic>k</italic>
<sup>2</sup>: roughly speaking, lowering it reduces crosstalk, but also increases losses. Changing the spacing of MRRs will have an impact on the overall footprint of the device. The number of components for parallelism in turn impacts the amount of added time that may be incurred if operations must be performed sequentially, in case the data size exceeds the capacity of a single optical element. One way to normalize for these effects is to assess the efficiency of the WDM usage in terms of energy per wavelength utilized [<xref ref-type="bibr" rid="B143">143</xref>].</p>
<p>In practice, it can sometimes be more efficient to use hybrid methods that offload some network tasks to standard electronic implementations, in which case energy consumption and speed limitations are incurred in optoelectronic conversion. The added energy expense can come from the receiver stages that follow detection, which may consist of amplification, sampling, and quantization [<xref ref-type="bibr" rid="B129">129</xref>]. Many accelerators apply nonlinear layers in the electronic domain, and some even combine photonic multiplication with electrical addition [<xref ref-type="bibr" rid="B144">144</xref>, <xref ref-type="bibr" rid="B145">145</xref>], especially when network weights or activations are reduced to one-bit representations. Optoelectronic conversions can introduce speed bottlenecks not only through DAC/ADC conversion steps but also by reverting to a dependence on electronic clock rate for sequential operations.</p>
</sec>
</sec>
<sec id="s4">
<title>4 Integrated photonic deep learning accelerators</title>
<p>In this section, we discuss examples of integrated circuit PDLAs which explore the challenge of mapping deep learning onto photonic hardware, showing comparative advantages and tradeoffs in various approaches. We group the accelerators on an application level according to important deep learning use cases: convolutional networks; linear models and sequence processing; and real-time or edge computing applications. These examples implement popular existing neural network architectures, which can facilitate nearer-term adoption. We provide two tables to aggregate main operating principles (<xref ref-type="table" rid="T1">Table 1</xref>), and summarize features and figures-of-merit (<xref ref-type="table" rid="T2">Table 2</xref>). <xref ref-type="fig" rid="F6">Figure 6</xref> shows the high-level application categories.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>High-level properties of the accelerators featured in the review.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Accelerator (Year)</th>
<th align="left">NN types</th>
<th align="left">Analog vs. digital</th>
<th align="left">All-optical vs hybrid</th>
<th align="left">Main photonic components</th>
<th align="left">Optical nonlinearity (implementation)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">ADEPT [<xref ref-type="bibr" rid="B168">168</xref>] (2021)</td>
<td align="left">Linear, CNN,Transformer</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MZI</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">Albireo [<xref ref-type="bibr" rid="B143">143</xref>] (2021)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MRR accumulation, MZM multiplication</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">Ascend [<xref ref-type="bibr" rid="B171">171</xref>] (2022)</td>
<td align="left">Linear, CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MRR weight banks</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">Bayesian [<xref ref-type="bibr" rid="B166">166</xref>] (2022)</td>
<td align="left">Bayesian NN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MZI mesh</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">Bitwise [<xref ref-type="bibr" rid="B155">155</xref>] (2021)</td>
<td align="left">CNN</td>
<td align="left">Digital</td>
<td align="left">Both versions</td>
<td align="left">MZI (optical accumulate), MRR (optical AND)</td>
<td align="left">Tanh (piecewise-linear approx. w/bit mapping)</td>
</tr>
<tr>
<td align="left">BPLight-CNN [<xref ref-type="bibr" rid="B192">192</xref>] (2021)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MRR weight banks</td>
<td align="left">ReLU (SOA), maxpool (optical comparators)</td>
</tr>
<tr>
<td align="left">ConvLight [<xref ref-type="bibr" rid="B146">146</xref>] (2017)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MRR weight banks</td>
<td align="left">ReLU (SOA), maxpool (optical comparators)</td>
</tr>
<tr>
<td align="left">CrossLight [<xref ref-type="bibr" rid="B149">149</xref>] (2021)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MRR weight banks, hybrid tuning</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">DNNARA [<xref ref-type="bibr" rid="B160">160</xref>] (2020)</td>
<td align="left">CNN</td>
<td align="left">Digital</td>
<td align="left">All-optical</td>
<td align="left">MRR for WDM, hybrid plasmonic- photonic (HPP) 2 &#xd7; 2 switch</td>
<td align="left">Sigmoid (RNS approx.)</td>
</tr>
<tr>
<td align="left">DNNARA-E [<xref ref-type="bibr" rid="B145">145</xref>] (2022)</td>
<td align="left">CNN</td>
<td align="left">Digital</td>
<td align="left">Hybrid</td>
<td align="left">MRR for WDM, hybrid plasmonic- photonic (HPP) 2 &#xd7; 2 switch</td>
<td align="left">Sigmoid, ReLU, maxpool (RNS approx.)</td>
</tr>
<tr>
<td align="left">FICONN [<xref ref-type="bibr" rid="B194">194</xref>] (2023)</td>
<td align="left">Linear</td>
<td align="left">Analog</td>
<td align="left">All-optical</td>
<td align="left">MZI mesh MVM</td>
<td align="left">ReLU (MZI phase shift)</td>
</tr>
<tr>
<td align="left">HolyLight [<xref ref-type="bibr" rid="B45">45</xref>] (2019)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">Microdisks</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">HQNNA [<xref ref-type="bibr" rid="B159">159</xref>] (2022)</td>
<td align="left">CNN</td>
<td align="left">Digital</td>
<td align="left">Hybrid</td>
<td align="left">MRR banks, hybrid tuning, VCSEL arrays</td>
<td align="left">Sigmoid (SOA)</td>
</tr>
<tr>
<td align="left">LightBulb [<xref ref-type="bibr" rid="B156">156</xref>] (2020)</td>
<td align="left">CNN</td>
<td align="left">Hybrid</td>
<td align="left">Hybrid</td>
<td align="left">Racetrack memory, microdisk XNOR gate, PCM-based ADC</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">LiteCON [<xref ref-type="bibr" rid="B193">193</xref>] (2022)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">All-optical</td>
<td align="left">Microdisk multiplication, crossbar array</td>
<td align="left">ReLU (SOA), maxpool (optical comparator)</td>
</tr>
<tr>
<td align="left">Mindreading [<xref ref-type="bibr" rid="B176">176</xref>] (2020)</td>
<td align="left">Linear, RNN, CNN</td>
<td align="left">Digital</td>
<td align="left">Hybrid</td>
<td align="left">Microdisk adders and shifters</td>
<td align="left">Logistic, tanh, ReLU (quantized approx.)</td>
</tr>
<tr>
<td align="left">Netcast [<xref ref-type="bibr" rid="B174">174</xref>] (2022)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MZM</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">PCNNA [<xref ref-type="bibr" rid="B148">148</xref>] (2018)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MRR weight banks</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">PIXEL [<xref ref-type="bibr" rid="B144">144</xref>] (2020)</td>
<td align="left">CNN</td>
<td align="left">Digital</td>
<td align="left">Both versions</td>
<td align="left">MZI (optical accumulate), MRR (optical AND), RF memory</td>
<td align="left">Tanh (piecewise-linear approx. w/bit mapping)</td>
</tr>
<tr>
<td align="left">RecLight [<xref ref-type="bibr" rid="B164">164</xref>] (2022)</td>
<td align="left">RNN</td>
<td align="left">Analog</td>
<td align="left">All-optical</td>
<td align="left">MRR banks, VCSEL arrays, memristors, hybrid tuning</td>
<td align="left">Sigmoid (SOA)</td>
</tr>
<tr>
<td align="left">ROBIN [<xref ref-type="bibr" rid="B158">158</xref>] (2021)</td>
<td align="left">CNN</td>
<td align="left">Digital</td>
<td align="left">All-optical</td>
<td align="left">MRR banks, hybrid tuning, VCSEL arrays</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">SONIC [<xref ref-type="bibr" rid="B150">150</xref>] (2022)</td>
<td align="left">CNN</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MRR banks, hybrid tuning, VCSEL arrays</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">Tiled MM [<xref ref-type="bibr" rid="B175">175</xref>] (2023)</td>
<td align="left">Linear</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MZI mesh, coherent crossbar</td>
<td align="left">N/A</td>
</tr>
<tr>
<td align="left">TRON [<xref ref-type="bibr" rid="B163">163</xref>] (2023)</td>
<td align="left">Transformer</td>
<td align="left">Analog</td>
<td align="left">Hybrid</td>
<td align="left">MRR banks, hybrid tuning, VCSEL arrays</td>
<td align="left">GELU (SOA)</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Features, figures-of-merit, and applications of accelerators. We reproduce metrics in the form reported by the paper, as not all accelerators report consistent figures-of-merit. Approximate values are indicated by &#x201c;&#x223c;&#x201d; where only relativevalues were reported, or were only reported visually in a plot. &#x201c;&#x2014;&#x201d; indicates that a value was not directly reported in the paper. Acronyms: GOPS &#x003D; giga operations/second; IPS &#x003D; inferences/second; FPS &#x003D; frames/second; MVM &#x003D; matrix-vector multiplication; EPB &#x003D; energy per bit.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Accelerator (Year)</th>
<th align="left">Features</th>
<th align="left">Figures-of-merit (reported)&#x2a;</th>
<th align="left">Network architecture</th>
<th align="left">Task (accuracy)&#x2a;&#x2a;</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">ADEPT [<xref ref-type="bibr" rid="B168">168</xref>] (2021)</td>
<td align="left">focuses on accelerated GEMM; can be applied in multiple network types</td>
<td align="left">- 10.59&#xa0;IPS/W/mm<sup>2</sup>
<break/>- 7,476.78&#xa0;IPS/W<break/>- 217, 201 IPS</td>
<td align="left">ResNet-50, BERT-large, RNN-T</td>
<td align="left">within 1% of benchmarks</td>
</tr>
<tr>
<td align="left">Albireo [<xref ref-type="bibr" rid="B143">143</xref>] (2021)</td>
<td align="left">provides analysis of bit precision; distributes across locally connected groups for added parallelism</td>
<td align="left">- 124.6&#xa0;mm<sup>2</sup> area<break/>- 395 GOPS/mm<sup>2</sup>
<break/>- 17.7 GOPS/W/mm<sup>2</sup>
</td>
<td align="left">VGG-16, ResNet18, MobileNet, AlexNet</td>
<td align="left">-</td>
</tr>
<tr>
<td align="left">Ascend [<xref ref-type="bibr" rid="B171">171</xref>] (2022)</td>
<td align="left">uses photonics for chip interconnects</td>
<td align="left">770&#xa0;mm<sup>2</sup> area (24.07/chiplet)</td>
<td align="left">VGG-16, ResNet-50, DenseNet-201, EfficientNet-B7</td>
<td align="left">-</td>
</tr>
<tr>
<td align="left">Bayesian [<xref ref-type="bibr" rid="B162">162</xref>] (2022)</td>
<td align="left">implements network pruning; provides uncertainty characterization</td>
<td align="left">0.5&#xa0;W power</td>
<td align="left">Custom</td>
<td align="left">MNIST (&#x223c;81%)</td>
</tr>
<tr>
<td align="left">Bitwise [<xref ref-type="bibr" rid="B155">155</xref>] (2021)</td>
<td align="left">bit-level parallelism; circulant matrix formulation for bitwise MVM</td>
<td align="left" style=";">&#x223c;1,000&#xa0;mm<sup>2</sup> area (OOE)<break/>&#x223c;0.1&#xa0;mm<sup>2</sup> area (OEE)<break/>&#x223c;100&#xa0;Js energy-delay product</td>
<td align="left">AlexNet, ZFNet, ResNet-34, VGG-16, GoogleNet</td>
<td align="left">ImageNet (&#x2212;)</td>
</tr>
<tr>
<td align="left">BPLight-CNN [<xref ref-type="bibr" rid="B192">192</xref>] (2021)</td>
<td align="left">supports training</td>
<td align="left">- 90,985 GOPS (inference)<break/>- 44,030&#xa0;GOPS/mm<sup>2</sup> (inference)<break/>- 9,327.5&#xa0;GOPS/W</td>
<td align="left">VGG, LeNet</td>
<td align="left">MNIST (95%)</td>
</tr>
<tr>
<td align="left">ConvLight [<xref ref-type="bibr" rid="B146">146</xref>] (2017)</td>
<td align="left">early example of end-to-end network</td>
<td align="left">- 15,000 GOPS/W<break/>- 20,000 GOPS/mm<sup>2</sup>
<break/>- 1.8&#xa0;mm<sup>2</sup> area (weight banks)</td>
<td align="left">VGG</td>
<td align="left">MNIST (94%)</td>
</tr>
<tr>
<td align="left">CrossLight [<xref ref-type="bibr" rid="B149">149</xref>] (2021)</td>
<td align="left">designs for robustness to fabrication and runtime variations</td>
<td align="left">- 28.78&#xa0;pJ/bit<break/>- 52.59 kFPS/W<break/>- 0.9&#xa0;mm<sup>2</sup> area</td>
<td align="left">LeNet, custom</td>
<td align="left">Sign-MNIST (&#x223c;90%) STL10 (&#x223c;70%) CIFAR10 (&#x223c;75%) Omniglot (&#x223c;75%)</td>
</tr>
<tr>
<td align="left">DNNARA [<xref ref-type="bibr" rid="B160">160</xref>] (2020)</td>
<td align="left">applies residue arithmetic MVM</td>
<td align="left">- 12.6 GOPS/mm<sup>2</sup>/W<break/>- 55.64&#xa0;mm<sup>2</sup> area</td>
<td align="left">LeNet, VGG, DeepFace, ResNet</td>
<td align="left">-</td>
</tr>
<tr>
<td align="left">DNNARA-E [<xref ref-type="bibr" rid="B145">145</xref>] (2022)</td>
<td align="left">applies residue arithmetic MVM<break/>up to 80x speedup over GPU</td>
<td align="left">- 0.39 TOPS/mm<sup>2</sup>
<break/>- 3.22 TOPS/W<break/>- 24.91 GOPS/mm<sup>2</sup>
<break/>- 124.78&#xa0;mm<sup>2</sup> area</td>
<td align="left">LeNet, VGG, DeepFace, ResNet</td>
<td align="left">-</td>
</tr>
<tr>
<td align="left">FICONN [<xref ref-type="bibr" rid="B194">194</xref>] (2023)</td>
<td align="left">supports training<break/>experimentally validated</td>
<td align="left">- 34.2&#xa0;mm<sup>2</sup> area<break/>- 0.53&#xa0;TOPS<break/>- 9.8&#xa0;pJ/OP</td>
<td align="left">Custom</td>
<td align="left">vowel classification (92.7%)</td>
</tr>
<tr>
<td align="left">HolyLight [<xref ref-type="bibr" rid="B45">45</xref>] (2019)</td>
<td align="left">accelerates power-of-two quantized (P2Q) CNNs; achieves equivalent accuracy to electronic implementation</td>
<td align="left">- 280.42 (M version), 22.46 (A version)&#xa0;mm<sup>2</sup> area<break/>- &#x223c;10<sup>3</sup> (M), &#x223c;10<sup>5</sup> (A) FPS/W<break/>- &#x223c;10<sup>5</sup> (M), &#x223c;10<sup>6</sup> (A) FPS</td>
<td align="left">LeNet, ResNet-18, AlexNet</td>
<td align="left">MNIST (98.9% LeNet-5) ImageNet (79.4% AlexNet, 88.6% ResNet-18)</td>
</tr>
<tr>
<td align="left">HQNNA [<xref ref-type="bibr" rid="B159">159</xref>] (2022)</td>
<td align="left">applies both WDM and TDM; supports different precision among layers</td>
<td align="left" style=";">- 57.5&#xa0;W power<break/>- &#x223c;10<sup>14</sup> GOPS/EPB</td>
<td align="left">AlexNet, ResNet-20, custom</td>
<td align="left">CIFAR10 (76.4% AlexNet, 79.7% ResNet) SVHN (87.9% custom)</td>
</tr>
<tr>
<td align="left">LightBulb [<xref ref-type="bibr" rid="B156">156</xref>] (2020)</td>
<td align="left">uses binarized CNN weights; photonic implementations of XNOR, ADC, and I/O</td>
<td align="left" style=";">- 24.05&#xa0;mm<sup>2</sup> area<break/>- 65.83&#xa0;W<break/>- &#x223c;10<sup>3</sup> FPS/W<break/>- &#x223c;10<sup>5</sup> FPS</td>
<td align="left">MobileNet, ShuffleNet, ResNet</td>
<td align="left">ImageNet (MobileNet 91.4%, ShuffleNet 87.3%, ResNet 87.9%)</td>
</tr>
<tr>
<td align="left">LiteCON [<xref ref-type="bibr" rid="B193">193</xref>] (2022)</td>
<td align="left">supports training<break/>292x potential speedup over GPU</td>
<td align="left">- 90,853 (train), 98,958 (test) GOPS<break/>- 1,132.85&#xa0;GOPS/W (avg.)</td>
<td align="left">VGG-Net, LeNet</td>
<td align="left">ImageNet (98%)</td>
</tr>
<tr>
<td align="left">Mindreading [<xref ref-type="bibr" rid="B176">176</xref>] (2020)</td>
<td align="left">real-time EEG analysis application<break/>minimizes power budget</td>
<td align="left">- 21.55&#xa0;W<break/>- 0.08041&#xa0;mm<sup>2</sup> area<break/>- 1000&#xa0;IPS/W</td>
<td align="left">EEG-Net</td>
<td align="left">EEG classification (97.6%)</td>
</tr>
<tr>
<td align="left">Netcast [<xref ref-type="bibr" rid="B174">174</xref>] (2022)</td>
<td align="left">edge compute application<break/>experimentally validated</td>
<td align="left">&#x3c; 1 photon/MAC (effective)</td>
<td align="left">Custom</td>
<td align="left">MNIST (98.8%)</td>
</tr>
<tr>
<td align="left">PCNNA [<xref ref-type="bibr" rid="B148">148</xref>] (2018)</td>
<td align="left">MRR bank &#x2b; BW protocol BW proof of concept</td>
<td align="left">2.2&#xa0;mm<sup>2</sup> area (weight banks)</td>
<td align="left">Custom</td>
<td align="left">-</td>
</tr>
<tr>
<td align="left">PIXEL [<xref ref-type="bibr" rid="B144">144</xref>] (2020)</td>
<td align="left">applies serial-parallel multiplication</td>
<td align="left">0.1 (OE), 100 (OO) &#x3bc; m<sup>2</sup> (MAC unit) 1503 (OE), 1044 (OO)&#xa0;mJ (ResNet)</td>
<td align="left">AlexNet, VGG, ResNet-34</td>
<td align="left">-</td>
</tr>
<tr>
<td align="left">RecLight [<xref ref-type="bibr" rid="B164">164</xref>] (2022)</td>
<td align="left">first non-coherent photonic RNN accelerator</td>
<td align="left">- 10<sup>4</sup> GOPS<break/>- 10<sup>-9</sup> J/bit</td>
<td align="left">Custom</td>
<td align="left">Weather prediction (0.5650 MAE) IMDB analysis (76.8%) Penn Treebank (65.78 perplexity)</td>
</tr>
<tr>
<td align="left">ROBIN [<xref ref-type="bibr" rid="B158">158</xref>] (2021)</td>
<td align="left">heterogeneous MR precision<break/>performs noise injection analysis</td>
<td align="left">- &#x223c;1.5e6 (EO)<break/>- &#x223c;3.25e6 (PO) FPS<break/>- &#x223c;10<sup>5</sup> FPS/W</td>
<td align="left">Custom</td>
<td align="left">Sign MNIST (&#x223c;92%) CIFAR10 (&#x223c;92.5%) STL10 (&#x223c;91%) SVHN (&#x223c;97%)</td>
</tr>
<tr>
<td align="left">SONIC [<xref ref-type="bibr" rid="B150">150</xref>] (2022)</td>
<td align="left">designed around network compression methods</td>
<td align="left">- &#x223c;10<sup>5</sup> FPS/W<break/>- &#x223c;10<sup>&#x2212;11</sup> J/bit</td>
<td align="left">Custom</td>
<td align="left">MNIST (92.89%) CIFAR10 (86.86%) STL-10 (75.2%) SVHN (95%)</td>
</tr>
<tr>
<td align="left">Tiled MM [<xref ref-type="bibr" rid="B175">175</xref>] (2023)</td>
<td align="left">focus on linear operations; experimentally validated</td>
<td align="left">- 0.12 TMACs/mm<sup>2</sup>
<break/>- 0.816&#xa0;mm<sup>2</sup> area</td>
<td align="left">Custom</td>
<td align="left">detect DDoS attacks (63.6% Cohen&#x2019;s kappa score)</td>
</tr>
<tr>
<td align="left">TRON [<xref ref-type="bibr" rid="B163">163</xref>] (2023)</td>
<td align="left">highly relevant architecture and application</td>
<td align="left">- 1e6 GOPS<break/>- 1e-10 J/bit</td>
<td align="left">Transformer, BERT, ViT, Albert</td>
<td align="left">TED translate (70.4%) BERT IMDB analysis (85.8%) Albert IMDB analysis (88.7%) ViT-base ImageNet (98.0%)</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Areas of application which can be accelerated by PDLAs.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g006.tif"/>
</fig>
<sec id="s4-1">
<title>4.1 Focus on CNNs</title>
<p>A prominent approach in photonic accelerators for deep learning is focused on implementing convolutional neural networks (CNNs) for fast photonic inference on computer vision tasks. Many convolution accelerators are based on WDM and resistive memory, which are implemented through configurations of components such as ring resonators, modulators, and interferometers. The WDM parallelism can be applied in an analog manner, or in a digital manner acting on different bits in parallel.</p>
<p>An early entry into photonic CNN accelerators was ConvLight, introduced by Dang et al. [<xref ref-type="bibr" rid="B146">146</xref>]. ConvLight implements an end-to-end architecture, with feature extraction blocks applying memristive convolution, semiconductor-optical-amplifier (SOA) ReLU activation, cascaded optical comparators for max pooling, and finally a memristive linear layer. The convolution unit comprises a WDM waveguide, a Weight Resistor Array (WRA) based on memristors, a Ring Modulator Array (RMA), and an SRAM buffer (SB). Weight values are stored in memristor conductance, which can be dynamically adjusted by applying an external current flux. Each weight bank in a weight resistor array consists of 9 memristors, representing a (3 &#xd7; 3) convolution filter. The output currents from these memristors are accumulated and fed into a modulator, where SOAs modulate the values for the element-wise ReLU activation. Post modulation, the modes are dropped from the WDM demux using a decoupler, and each isolated lightwave is then directed to the subsequent layer. Successive feature extraction units are joined by electronic interface layers. Finally, the accumulated current from each memristor bank is digitized for an output value. When compared to the FPGA-based Caffeine accelerator [<xref ref-type="bibr" rid="B7">7</xref>] and memristor crossbar-based ISAAC accelerator [<xref ref-type="bibr" rid="B136">136</xref>], ConvLight showed 250&#xd7; and 28&#xd7; higher CE, respectively. These comparisons were based on training and inference tasks executed on four versions of the VGG [<xref ref-type="bibr" rid="B124">124</xref>] model applied to the MNIST dataset [<xref ref-type="bibr" rid="B147">147</xref>].</p>
<p>Notably, ConvLight uses one memristor for each weight, making its footprint scale with the number of network parameters. Mehrabian et al. [<xref ref-type="bibr" rid="B148">148</xref>] later introduced PCNNA, a proof-of-concept analog design which presents improved usage of parallelism with MRR weight banks structured based on the broadcast-and-weight (BW) protocol. <xref ref-type="fig" rid="F7">Figure 7</xref> depicts the basic formulation of the protocol. PCNNA makes use of the fact that the same kernel values of the layer are applied to all windows of the input, and that iterating over all the windows not costly in a photonic implementation, in comparison with conventional hardware. They use microrings only for the kernel receptive field of size <italic>k</italic>, multiplied by the number of kernels for the output depth. Given that the kernels share the same receptive field of the input, they can be executed in parallel. <xref ref-type="fig" rid="F8">Figure 8</xref> shows the difference in their approach. Overall, this reduces both the number of wavelengths required to represent the input feature map, and the number of microrings needed at the following layer for demultiplexing. They show that in execution time, the iteration over receptive fields fits within a single slow clock cycle.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>An MRR bank-based broadcast-and-weight protocol. A bundled wavelength is propagated through an MRR bank as it enters. Through the tuning of corresponding rings, each bank weights each wavelength. Photodiodes create photocurrents by adding all wavelengths together. Photo-currents modulate light waves of wavelength <italic>&#x3bb;</italic>
<sub>
<italic>m</italic>
</sub>. Multiplexing of all laser beams is used to broadcast the beams to the next layer. Reproduced with permission, from [<xref ref-type="bibr" rid="B148">148</xref>] Mehrabian et al. 2018 &#xa9; IEEE.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g007.tif"/>
</fig>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>Illustration of MRR bank use in convolution: a 16 &#xd7; 16 input feature map with 5 kernels of 3 &#xd7; 3 is implemented in <bold>(A)</bold> using one ring per input wavelength, whereas <bold>(B)</bold> uses only one ring per distinct kernel value required to cover the receptive field, which results in fewer required MRRs. Reproduced with permission, from [<xref ref-type="bibr" rid="B148">148</xref>], Mehrabian et al. 2018 &#xa9; IEEE.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g008.tif"/>
</fig>
<p>Otherwise, the photonic multiplication flow takes place as usual: a waveguide is employed as a transmission line to broadcast multiplexed wavelengths to the next layer, such that each neuron in the destination layer receives all incoming wavelengths. The amplitude of each wavelength at the output is determined by a weighting function corresponding to the incident power and biasing potential of the MRR. Following multiplication, a photodiode integrates all incoming wavelengths, generating an aggregate photocurrent to implement the accumulation operation. With this design, a representative layer of a network such as AlexNet [<xref ref-type="bibr" rid="B4">4</xref>] can be evaluated 3 orders of magnitude faster than electronic computation, even including the time cost caused by electronic I/O.</p>
<p>In contrast, it is also possible to implement parallelism along the receptive field dimension. Shiflett et al. take this approach in their Albireo accelerator [<xref ref-type="bibr" rid="B143">143</xref>]. In their construction, computation is performed concurrently on multiple receptive fields of the input. The Photonic Locally Connected Units (PLCUs) of Albireo contain a grid of MRRs, where the input dimension is the number of kernel elements represented by MZMs, and the output dimension is the number of receptive fields, each transmitted to an output photodetector. Each PLCU processes a single channel of the convolution, simultaneously computing on all receptive fields. However, to maintain sufficient analog precision, the maximum number of wavelengths for each PLCU is restricted, so to process more fields simultaneously, multiple PLCUs are clustered in PLC groups (PLCGs). Overall, each PLCG implements a single kernel of the layer, acting on the same input volume in parallel, which is broadcast to all PLCGs at the same time. This distributed structure also gives Albireo the ability to implement depth-wise separable convolution layers, which are often used in practice. Albireo illustrates how parallelism is constrained by the number of possible wavelengths, informing design choices based on the expected dimensions of the kernel size, number of kernels, and number of receptive fields in the input.</p>
<p>Optimizations can also be made to balance the overall configuration of the weight banks. Non-coherent architectures are highly susceptible to process variations, as well as runtime variations induced by heat and environmental factors. For instance, increasing the length of the waveguide hosting the MR banks increases the total optical signal propagation, modulation, and losses, which in turn increases the laser power required for optical signals to be detected error-free; crosstalk noise can also substantially deteriorate weight resolution [<xref ref-type="bibr" rid="B149">149</xref>]. In their CrossLight accelerator, Sunny et al. [<xref ref-type="bibr" rid="B149">149</xref>] perform device-level optimizations to improve robustness. They make adaptations such as hybrid thermo-optic and electro-optic tuning to compensate for thermal crosstalk, and determine an optimal number of MRs per wave bank which can still support 16-bit resolution. They take into consideration layout spacing, wavelength reuse within weight banks, and optical splitter losses. They report that the final optimized configuration has 9.5&#xd7; lower energy-per-bit and 15.9&#xd7; higher performance-per-watt over other photonic accelerators.</p>
<p>Sunny et al. introduce another approach to increasing efficiency with SONIC [<xref ref-type="bibr" rid="B150">150</xref>], an accelerator architecture optimized for networks that have been compressed using techniques developed in deep learning practice [<xref ref-type="bibr" rid="B151">151</xref>]: <xref ref-type="fig" rid="F9">Figure 9</xref> depicts the SONIC accelerator architecture.. The first compression technique is to apply sparsity-aware training to induce layer-wise sparsity [<xref ref-type="bibr" rid="B152">152</xref>]. In the accelerator, sparse and dense vectors can be buffered separately, and the sparse input path uses power gating to prevent VCSELs from being driven for a zero element. The second technique is clustering model weights post-training to restrict to a fixed number of unique weights. This assumption allows for lower resolution requirements in DAC conversion. In SONIC, sparse vector weights can be reduced to 6 bits, while dense activation values are kept at 16 bits. This separation of pathways is reflected in the overall architecture, as shown in <xref ref-type="fig" rid="F6">Figure 6</xref>. These adaptations allow SONIC to improve energy-per-bit 8.4&#xd7; and power efficiency 5.8&#xd7; over electronic accelerators.</p>
<fig id="F9" position="float">
<label>FIGURE 9</label>
<caption>
<p>An overview of the SONIC architecture, showing the distinct pathways of data which participates in either sparse or dense computations. Reproduced with permission, from [<xref ref-type="bibr" rid="B150">150</xref>], Sunny et al. 2022 &#x00A9; IEEE.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g009.tif"/>
</fig>
<p>In contrast with analog methods, some accelerators operate in a digital paradigm, using photonic parallelism for concurrent bitwise and logical tasks. An early example within this domain is the HolyLight accelerator, as introduced by Liu et al [<xref ref-type="bibr" rid="B45">45</xref>], which is designed to accelerate power-of-2 quantized (P2Q) CNNs [<xref ref-type="bibr" rid="B153">153</xref>]. The device incorporates matrix-vector multipliers (MVMs), and a 16-bit ripple-carry adder constructed from full adders using microdisks, alongside P2Q-CNN inference units. This system uses digital electronics to compute the generate and propagate values from the output of each full adder, while the photonic accelerator calculates the sum and carry operations. Two variations of this architecture were developed to explore different aspects of computational efficiency, including the maximum speed of MRR operation, as well as considerations related to noise and signal degradation. HolyLight-M incorporates digital-to-analog converters (DACs) and analog-to-digital converters (ADCs) for the transition between digital values and optical signals. HolyLight-A integrates multiple photonic shifters and adders, connected through a shared bus system. Both variants of HolyLight demonstrate a 5&#xd7; improvement in power efficiency compared to traditional GPU, CPU, and TPU architectures. <xref ref-type="fig" rid="F10">Figure 10</xref> shows the overall flow of the accelerator design.</p>
<fig id="F10" position="float">
<label>FIGURE 10</label>
<caption>
<p>Diagram of the Holylight accelerator architecture [<xref ref-type="bibr" rid="B42">45</xref>]. <bold>(A)</bold> shows the overall chip node, which consists of multiple connected tiles <bold>(B)</bold>. Tiles contain Photonic Processing Units (PPUs). <bold>(C)</bold> is the PPU structure of the HolyLight-M variant, and <bold>(D)</bold> is the HolyLight-A PPU. Reproduced with permission from [<xref ref-type="bibr" rid="B42">45</xref>], Liu et al. &#xa9; EDAA.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g010.tif"/>
</fig>
<p>The PIXEL accelerator of Shiflett et al. [<xref ref-type="bibr" rid="B144">144</xref>] is a photonic accelerator that uses a combination of MRRs for bitwise logic operations, and MZMs for accumulation. Mathematically, PIXEL is modeled after the Stripes (STR) [<xref ref-type="bibr" rid="B154">154</xref>] formulation of accelerated neural networks through serial-parallel multiplication. In this method, the computational time is linear in the length of the serial input, which is the bit precision of a given network layer. The authors present efficient photonic implementations, with one hybrid Optical-Electrical (OE) version that multiplies in the optical domain and then accumulates in the electrical domain, and a fully Optical-Optical (OO) version for both multiplying and accumulating in the optical domain. PIXEL&#x2019;s OMAC units use radio frequency memory for storing filter weights in addition to the MAC unit.</p>
<p>In PIXEL, each MZM accumulates a single wavelength, which increases the number of MZMs in their design, reducing area efficiency. Later, Shiflett et al. [<xref ref-type="bibr" rid="B155">155</xref>] advanced on the PIXEL design to improve the usage of WDM by implementing parallelism in bit-wise operations. In this design, the bitwise matrix multiplication uses a circulant matrix formulation to take advantage of broadcasting a single bit value to multiple processing elements (PE). The authors again present two versions with different accumulation implementations. In both cases, MRRs are used to implement a bitwise AND operation. The first version then applies electronic processing for summation (O-E-E), while the other uses MZIs for accumulation, with a final electrical summation (O-O-E). The comparison with an all-electronic version of the accelerator shows that the EDP of the O-O-E implementation is 33.1% lower, and its speed is 79.4% faster.</p>
<p>Many accelerators based on logical operations rely on ripple-carry adders and SRAMs, both of which can limit the frequency and inference throughput of the accelerator when trying to replicate higher bit precision, due to the adder&#x2019;s long critical path and the SRAM&#x2019;s access latency. Zokaee et al. [<xref ref-type="bibr" rid="B156">156</xref>] take a distinct approach to address this problem by processing <italic>binarized</italic> CNNs rather than CNNs with floating point weights. Their accelerator, LightBulb, uses microdisks to implement XNOR gates and popcount operations, followed by a photonic phase-change memory (pPCM) implementation of ADC. It also reduces input/output latency by using photonic racetrack memory, to enable 50&#xa0;GHz operating speed. To replace floating-point MACs with XNORs and popcounts, LightBulb first binarizes the weights and activations of a CNN into linear combinations of (&#x2212;1, &#x002B; 1)<italic>s</italic>, allowing the MRR to take advantage of bit-wise parallelism. pPCMs then achieve an ADC step photonically by implementing a temporal binary search [<xref ref-type="bibr" rid="B157">157</xref>]. LightBulb compares favorably against state-of-the-art GPU, FPGA, ASIC, ReRAM, and photonic CNN accelerators when tested on binarized MobileNet, ShuffleNet, and ResNet architectures. Overall, LightBulb achieves its efficiency by using photonic components for logical operations, ADC, and data I/O, which are typically large sources of latency and energy overhead. LightBulb improves throughput 17&#xd7; to 173&#xd7; over prior optoelectronic accelerators and increases throughput per Watt by 17.5&#xd7; to 660 &#xd7;.</p>
<p>The ROBIN accelerator from Sunny et al. [<xref ref-type="bibr" rid="B158">158</xref>] also makes use of binarization, but uses only binarized weights, leaving activation function values at 4-bit precision. This is intended to mitigate loss of accuracy. To implement this, ROBIN uses heterogeneous MRRs with different precisions, within an overall BW-based design, with improved pipelining of interactions with the electronic control unit. ROBIN also implements photonic batch normalization and adds circuit- and device-level optimizations intended to account for the effects of process variations. They perform extensive optimization over device configurations to develop two versions, one optimized for FPS performance (ROBIN-PO), and the other for area and energy efficiency (ROBIN-EO). ROBIN-EO achieves approximately 4x lower energy-per-bit than electronic BNN accelerators, whereas ROBIN-PO shows roughly 3x better performance than electronic BNN accelerators.</p>
<p>Later, Sunny et al. also applied mixed precision to reduce memory requirements with their Heterogeneous Quantization Neural Network Accelerator (HQNNA) [<xref ref-type="bibr" rid="B159">159</xref>]. HQNNA uses non-coherent photonics based on both WDM and a novel Time Division Multiplexing (TDM) approach with bit-slicing. The matrix-vector multiplication unit (MVU) performs multiplication and accumulation optically by distributing bit slices across time steps, then using digital shift and adder circuits to produce the final output. Bits that interact in the same dot product are assigned the same wavelength for photonic multiplication and transmitted in one step, with the resulting value shifted and buffered digitally after ADC. This is repeated for the number of bits per slice. This results in performing multiple smaller products rather than a single large product, which improves efficiency given the low latency and energy consumption of photonic multiplication. It also allows for heterogeneous precision across layers. This MVU design is applied both in linear and convolutional layers. HQNNA shows 52.2&#xd7; and 3.59&#xd7; improvement in EPB over LightBulb and ROBIN, respectively.</p>
<p>Peng et al. introduced another numerical innovation with DNNARA [<xref ref-type="bibr" rid="B160">160</xref>], which combines WDM with a Residue Number System (RNS). With RNS, a number can be represented as pairwise coprime moduli. Because residue arithmetic is digit-irrelevant, results can be combined separately during the residue operation and ensembled at the end, representing addition by mappings in the arithmetic system. Every modulo digit has a single-bit output without repetition, enabling computation-in-network using one-hot encoding photonic routing. RNS can allow for optical components with shorter optical critical paths, and the use of one-hot encoding also facilitates fast switching between the electrical and optical domains. However, the implementation of sigmoid activation functions like logistic and hyperbolic tangent with RNS is difficult. As a result, logistic and tanh functions are approximated by their Taylor series, and they can be implemented as polynomials with adders and multipliers. In subsequent work, the authors introduced DNNARA-E [<xref ref-type="bibr" rid="B145">145</xref>], which substitutes DNNARA&#x2019;s optical adders with electrical adders for reduced area, improved power usage, and ReLU activation function implementation. Overall, this results in three times better throughput than the original DNNARA. With a similar power budget, DNNARA-E achieves on average 80x speedup over the NVIDIA Tesla V100 GPU.</p>
</sec>
<sec id="s4-2">
<title>4.2 Beyond convolution</title>
<p>While convolutional neural networks remain an essential area of deep learning, many other architectures are important in practice and contribute to overall deep learning inference usage. This includes architectures that power ChatGPT and other sequence-based tasks, which can be extremely inefficient to evaluate on standard hardware. In addition, many computer vision models are also replacing convolutions with linear layers, as in the Vision Transformer [<xref ref-type="bibr" rid="B161">161</xref>]. Recent accelerator designs have begun to address this shift.</p>
<p>Importantly, the Transformer architecture has risen to prominence both in its original context of natural language sequence processing [<xref ref-type="bibr" rid="B162">162</xref>], and more recently as a strong alternative in image tasks [<xref ref-type="bibr" rid="B161">161</xref>]. To adapt to this trend, Afifi et al. [<xref ref-type="bibr" rid="B163">163</xref>] introduced TRON, the first SiPh hardware accelerator for Vision Transformers (ViTs). TRON utilizes non-coherent SiPh circuits to replicate the Transformer architecture&#x2019;s feedforward and multi-head attention (MHA) units. The required matrix multiplications are performed with an MR weight bank, with a design that efficiently pipelines the operations to re-use intermediate results. The softmax operation is efficiently approximated in the electronic domain, making TRON a hybrid model. TRON also replicates the GELU activation similarly to the method in a standard architecture, scaling the output data vector using an MR, applying a sigmoid function, and then applying MR multiplication again between this output and the data vector. MR units also implement normalization layers, and residual connections are performed through coherent summation. Depending on the application, Transformers may perform encoding only, or both encoding and decoding. TRON is structured so that decoder blocks can re-use the VCSEL arrays which drive input to the MHA unit. This reuse also introduces efficiency by reducing laser power consumption and crosstalk between channels. <xref ref-type="fig" rid="F11">Figure 11</xref> illustrates this overall structure. Software optimization techniques can also be applied to further reduce the memory footprint of the Transformer for additional performance improvement. TRON is simulated for popular Transformer-based models including BERT [<xref ref-type="bibr" rid="B162">162</xref>] and ViT. When compared against state-of-the-art GPU and FPGA accelerators, TRON shows 262&#xd7; better GOPs than general GPU benchmarks, and 55&#xd7; improvement over FPGA. It also improves energy-per-bit by 4,231&#xd7; over GPU, and 8&#xd7; for FPGA.</p>
<fig id="F11" position="float">
<label>FIGURE 11</label>
<caption>
<p>Overview of the TRON accelerator architecture, which replicates the multi-head attention and feedforward blocks of the Transformer architecture. Reproduced with permission, from [<xref ref-type="bibr" rid="B163">163</xref>], Afifi et al. 2023, &#x00A9; Association of Computing Machinery.</p>
</caption>
<graphic xlink:href="fphy-12-1369099-g011.tif"/>
</fig>
<p>Another essential class of neural networks is Recurrent Neural Networks (RNNs). Sunny et al. introduced a novel non-coherent photonic RNN accelerator called RecLight [<xref ref-type="bibr" rid="B164">164</xref>] which can accelerate NNs that consist of recurrent components, including Gated Recurrent Units (GRUs) and Long Short-Term Memory Networks (LSTMs) [<xref ref-type="bibr" rid="B165">165</xref>]. These architectures process sequence data by assigning a trainable &#x201c;hidden&#x201d; state to each sequence element. These weight matrices form connections across the sequence. Further, &#x201c;gating&#x201d; weights are optimized to either propagate information or suppress unnecessary pathways. To achieve the recurrent network structure, RecLight uses separate MAC units are used for input and hidden state weight matrices. RecLight achieves better parameter resolution by reducing thermal crosstalk, applying a hybrid tuning approach with both thermo-optic (TO) and electro-optic (EO) tuning. When compared with electronic RNN accelerators, RecLight improves energy-per-bit up to 1730&#xd7;, and has up to 2,631.6&#xd7; better throughput.</p>
<p>Sarantoglou et al. [<xref ref-type="bibr" rid="B166">166</xref>] explore the area of uncertainty quantification and Bayesian networks by introducing an accelerator with two innovative schemes: the first is the Bayesian regularized, aimed at reducing power consumption, and the second is the fully Bayesian, which offers insights into phase shifter sensitivity. Their approach focuses on the MNIST dataset [<xref ref-type="bibr" rid="B147">147</xref>] classification with 512 phase shifters, with their architecture similar to the one presented by Perez et al [<xref ref-type="bibr" rid="B167">167</xref>]. The system incorporates pre-characterization stages that monitor the variation between the applied current (I) and the induced phase shift (<italic>&#x3d5;</italic>). These pre-computational steps are designed to counter fabrication errors and inter-element crosstalk through passive offsets. Their findings demonstrate a significant reduction in the processing power required by the photonics integrated Circuit (PIC) without sacrificing classification accuracy. Moreover, the fully Bayesian scheme not only reduces energy consumption but also provides valuable data on phase shifter sensitivity. Consequently, this allows for the partial deactivation of phase actuators, substantially simplifying the driving system. The phase tuning process is based on an offline training scheme that takes into account uncertainty. Instead of defining optimum phase shifter values through training, a parametric Probability Distribution Function (PDF) is defined for each phase shifter and is optimized by updating variational parameters at every iteration. Aside from indicating the correct values for phase shifters, this Bayesian procedure also quantifies their robustness to phase deviation. Using this data, novel algorithms can be developed for adjusting and controlling photonic accelerators, which can further increase their robustness to noise and hence allow for increased scale.</p>
<p>In practice, many modern applications require greater flexibility than a straightforward translation of a network as a unit. To expand the use of photonic accelerators beyond cases that simply apply a fixed architecture, it is essential to develop devices with increased generality. For instance, Demirkiran et al. emphasize the relevance of linear acceleration and efficient matrix multiplication with their ADEPT accelerator [<xref ref-type="bibr" rid="B168">168</xref>]. ADEPT addresses important aspects of implementing linear layers, including the fact that most layer transforms are non-square, which can present a performance issue when multiplication and addition are combined in a single optical step. ADEPT favors an optoelectronic architecture combining optical general matrix-matrix multiplication (GEMM) operations with a digital electronic ASIC for nonlinear operations such as activations. In its pipeline, SVD and phase decomposition are performed on the original weights as an up-front digital operation. The design incorporates optimized buffering to minimize the speed bottleneck in optoelectronic transfer. They choose MZI components over MRR, citing their improved compatibility with electronic devices, which can facilitate the integration of the accelerator in practice. ADEPT can accommodate more modes as opposed to other accelerators, illustrating the benefit of a generalized design that can be compared to benchmarks beyond CNN applications. ADEPT shows competitive performance on benchmarks for ResNet-50 [<xref ref-type="bibr" rid="B169">169</xref>], BERT-large [<xref ref-type="bibr" rid="B162">162</xref>], and RNN-T [<xref ref-type="bibr" rid="B170">170</xref>]. They also report 2.5 &#xd7; better throughput per watt compared to state-of-the-art photonic accelerators.</p>
<p>Li et al. introduced the ASCEND accelerator [<xref ref-type="bibr" rid="B171">171</xref>], a chiplet-based system that utilizes the inherent low-latency characteristic of photonic interconnects to facilitate multi-chiplet broadcasting of data and weights within a neural network. This approach leverages the superior speed of optical communications over electrical interconnects [<xref ref-type="bibr" rid="B172">172</xref>, <xref ref-type="bibr" rid="B173">173</xref>]. By enabling chiplets to communicate seamlessly, ASCEND minimizes delays in mapping convolution layers both within and across chiplets. The accelerator&#x2019;s physical layout features columns and rows of local Processing Elements (PEs) organized into unit 2D arrays across chiplets. These PEs communicate with the Global Buffer (GLB) through a waveguide in a unicast manner, while broadcast communication from the GLB to each PE is also facilitated via a waveguide. This arrangement allows for the mapping of convolution layers at the granularity of the 2D PE array, ensuring efficient one-hop communication both within and between chips. ASCEND not only reduces energy consumption by 37% for DenseNet and 67% for ResNet-50 compared to chiplet-based accelerators with metallic interconnects but also achieves up to a 52% improvement in speed. This demonstrates the advantages in energy efficiency and processing speed gained by incorporating diverse photonic elements in accelerator architectures.</p>
</sec>
<sec id="s4-3">
<title>4.3 Alternative applications</title>
<p>Another approach focuses on matching the particular strengths of photonics to applications such as edge computing and real-time applications, as well as cases where initial analog-to-digital conversion of input data can be avoided, for a direct pipeline to optical inference. In this area, Sludds et al. introduced Netcast [<xref ref-type="bibr" rid="B174">174</xref>], a protocol that employs delocalized analog processing, performing efficient photonics inference using cloud-based smart transceivers to stream weight data to edge devices. This protocol is designed to facilitate the deployment of advanced neural network models on devices with strict power, processing, and memory constraints. Using wavelength division multiplexing (WDM), Netcast uses the optical spectrum for high-capacity data transmission by integrating cloud servers with smart transceivers that broadcast deep neural network weights. Optical matrix-vector multiplication is performed on-site in the edge devices equipped with broadband optical modulators. The weight matrix of one DNN layer is encoded on a time-frequency basis by the amplitude-modulated field. This is streamed to the client, which can modulate it using a broadband optical modulator to separate the wavelengths to N time-integrating detectors to produce the desired dot product. The Netcast design maximizes the number of MACs performed by every component in the client: in effect, this allows it to achieve an efficiency of less than one photon per MAC (0.1&#xa0;aJ/MAC). Netcast can be readily integrated into applications that operate on data streamed through existing commercial network switches. Through this method, milliwatt-class edge devices can compute at teraFLOPS rates, which were traditionally reserved for cloud computing infrastructures with much larger sizes and power consumption.</p>
<p>In another case, Giamougiannis et al. [<xref ref-type="bibr" rid="B175">175</xref>] introduced a coherent analog SiPho computing engine designed for fast optical Tiled Matrix Multiplication (TMM) at 50&#xa0;GHz. This accelerator incorporates Coherent Linear Neurons (COLNs) equipped with high-speed Silicon Germanium Electro-Absorption Modulators (EAMs) for both weight and input imprinting. The accelerator was deployed in a data center traffic inspection system for network security applications to highlight its practical capabilities in performing TMM. The photonic engine was experimentally tested for identifying Distributed Denial-of-Service (DDoS) attack patterns by classifying Reconnaissance Attacks (RAs). The size of the network is small: only 6 input features, one hidden layer of 8 neurons, and 2-neuron output. However, even this small classifier suffices to solve a practical use-case, demonstrating the advantage of integration into applications where replicating a large network size is not the primary aim.</p>
<p>Another interesting application is demonstrated by the ultra-low-power photonic MindReading accelerator by Lou et al. [<xref ref-type="bibr" rid="B176">176</xref>], intended for real-time processing of Electroencephalography (EEG) signals. The EEG device has a sampling rate of 128&#xa0;Hz, so MindReading seeks to minimize power consumption while matching this rate for inference. To do this, MindReading uses microdisks to perform energy- and area-efficient photonic shifting and adding operations. The accelerator utilizes logarithmic quantization applied to both weights and activations of convolution, recurrent, and fully connected layers. Floating point multiplication is replaced by addition and shift operations with a low bit width requirement so that precision can be reduced to 4 bits with minimal loss in accuracy. This accelerator replicates the structure of EEG-Net, which includes convolutional, fully-connected, and LSTM layers. The LSTM component requires sigmoid (tanh and logistic) activations, so MindReading uses a photonic unit for quantizing these functions. An eDRAM buffer is used for storing EEG signals as well as intermediate results generated by the Photonic Processing Unit (PPU). Then, by using photonic additions and shifters, the PPU computes binary logarithms and logarithmic accumulations for ULQ-quantized EEG-NET. MindReading reduces power consumption by 62.7% and increases throughput by 168.6% on average in comparison to existing accelerator counterparts for the same classification task. Overall, MindReading achieves approximately 1000 IPS (inferences per second) per Watt, whereas FPGA, CPU and GPU can reach less than 5 IPS per Watt.</p>
</sec>
</sec>
<sec id="s5">
<title>5 Discussion and research gaps</title>
<sec id="s5-1">
<title>5.1 Ongoing challenges</title>
<p>Despite the considerable advantages that photonic DL accelerators offer over their electronic counterparts, many challenges persist. In terms of design, the limited scale of PICs still restricts the numerical size of both the input vectors that can be loaded onto photonic hardware and the size and number of internal network layers. Challenges arise when scaling to larger matrices, due to the increasing number of phase shifters in MZI meshes that consume 15&#xa0;mW on average [<xref ref-type="bibr" rid="B99">99</xref>]. The power consumption required for large MAC operations would necessitate thousands of such phase shifters, which increases the cost of thermal management. As an alternative, NOEMS devices have been considered a suitable replacement due to their near-zero static power dissipation as they wiggle the waveguide back and forth [<xref ref-type="bibr" rid="B177">177</xref>]. For WDM systems, the in put supported is ultimately limited by the number of wavelength channels that can be multiplexed on a single waveguide bus. However, the number of neurons can be expanded with spectrum reuse techniques for the WDM schemes as reported in [<xref ref-type="bibr" rid="B129">129</xref>]. Another challenge to scaling is the implementation of caching memory subsystems, which becomes difficult when handling real workloads generating substantial intermediate data. To execute large-scale neural networks, electronic memories such as SRAM and DRAM can be integrated with optical video memory modules [<xref ref-type="bibr" rid="B178">178</xref>].</p>
<p>In the case of coherent approaches, scaling the network can be restricted by the number of required components. Shafiee et al. [<xref ref-type="bibr" rid="B179">179</xref>] have conducted an extensive comparison among the Reck, Clements, and Diamond designs to assess their comparative robustness to optical loss and crosstalk noise. This evaluation was carried out by measuring the degradation in inference accuracy with an increased mesh scale. Their work highlights the drawbacks of increasing scale primarily by increasing mesh size. However, advances in PPC design present alternatives where the number of elements in a mesh can be reduced without compromising computational capacity. For instance, in addition to the reduced footprint resulting from a different configuration of components, algorithmic improvements can enhance the fidelity in computing of photonic unitary operation, as shown by Yu and Park [<xref ref-type="bibr" rid="B180">180</xref>]. They build on the Clements design by introducing the &#x201c;pruning&#x201d; of redundant rotations in the computed operators. The design of Buddhiraju et al. [<xref ref-type="bibr" rid="B181">181</xref>] applies a resonator-based architecture utilizing the frequency synthetic dimension, to achieve <italic>O</italic>(<italic>N</italic>) scaling in footprint and gate numbers. They report a higher compute density than MZI meshes at approx. 10 TMACs per second per unit area (mm<sup>2</sup>), which is comparable with silicon crossbar designs. However, the values of N are restricted by the free spectral range of the sources and the device bandwidths. Recent work by Piao et al. [<xref ref-type="bibr" rid="B182">182</xref>] focuses on a method that exploits space-time duality for programmable photonic &#x201c;time&#x201d; circuits (PPTCs). PPTCs use coupled resonators which substitute spatial optical path length with field evolution in the time domain. This design achieves <italic>O</italic>(<italic>N</italic>) scaling in both spatial circuit footprint and the number of optical gates. Other important contributions have also been made to advance error correction mechanisms, mitigating fabrication-induced inaccuracies that compromise the performance of large-scale systems. Bandyopadhyay et al. [<xref ref-type="bibr" rid="B183">183</xref>] focuses on improving the fidelity of MZI and mesh-based unitary operations such that, at an application level, developers can assume the underlying hardware will maintain fidelity. Their proposed method involves deterministic correction of hardware errors within optical gates using local corrections. Overall, these examples present interesting possibilities for pushing scale boundaries for PPC-based accelerators.</p>
<p>In the case of MRR-based noncoherent approaches, scaling up can also present problems with increased power requirements and thermal accumulation. Such as nano-optoelectromechanical systems (NOEMS) [<xref ref-type="bibr" rid="B184">184</xref>] and liquid crystals on silicon LCOS [<xref ref-type="bibr" rid="B67">67</xref>], can notably improve energy efficiency due to the low voltage bias conditions. Efficient weight tuning is achievable with low-loss thermal phase shifters. Moreover, high speed, low voltage swing modulators (1-2&#xa0;Vpp) [<xref ref-type="bibr" rid="B185">185</xref>, <xref ref-type="bibr" rid="B186">186</xref>] promise improved energy efficiency by consuming less power on the CMOS driver and modulator [<xref ref-type="bibr" rid="B67">67</xref>]. Other improvements can be achieved using integrated photoconductive heaters [<xref ref-type="bibr" rid="B187">187</xref>] with resonant tuning over a wide dynamic range without the need for additional tuning mechanisms or additional electrical interfaces. Integrated with silicon photonics, lithium niobate, and barium titanate electro-optical modulators provide high-speed phase modulation and low operating voltage, making them extremely attractive for photonic accelerators.</p>
<p>In addition to such improvements in the photonic mechanisms, in order to make advances in practical adoption, it is essential to improve standardization in the reporting of design and performance statistics. Comparisons can be hard to make across the literature, as there is limited consistency and completeness in reporting metrics, in terms of hardware and network configurations as well as datasets. Some accelerators do not report inference accuracy, which is a critical statistic considered by deep learning practitioners. It is also crucial to consider that for photonic accelerators to be adopted, they must either integrate seamlessly with existing deep learning platforms such as PyTorch, or present similar user-friendly software libraries where application-level adaptations can be made.</p>
<p>Finally, it is important to note that most accelerators which have been practically implemented still focus on inference with offline training. While this is particularly useful in real-time applications requiring high-speed inference, or edge computing with resource constraints, in practice, the bulk of arithmetic intensity in deep learning is incurred during the training process. Attention has increasingly shifted toward designing accelerators that can execute online photonic training. Buckley et al. [<xref ref-type="bibr" rid="B188">188</xref>] provide a recent survey on the status of training capability in PDLAs. Hughes et al. introduced a theoretical treatment considering algorithmic aspects of training in optical platforms [<xref ref-type="bibr" rid="B72">72</xref>], and more recently this proposal has been realized experimentally with over 94% accuracy on the MNIST digit recognition task [<xref ref-type="bibr" rid="B189">189</xref>]. Free-space devices have been studied by Spall et al. in both hybrid [<xref ref-type="bibr" rid="B190">190</xref>] and all-optical [<xref ref-type="bibr" rid="B191">191</xref>] variants. Dang et al. have proposed extensions of their ConvLight design which can also accommodate training [<xref ref-type="bibr" rid="B192">192</xref>, <xref ref-type="bibr" rid="B193">193</xref>].</p>
<p>Bandyopadhyay et al. [<xref ref-type="bibr" rid="B194">194</xref>] have advanced this in the area of PIC by fabricating and testing an all-optical device that performs both inference and <italic>in situ</italic> training. Their Fully Integrated Coherent Optical Neural Network (FICONN) system incorporates Nonlinear Optical Function Units (NOFUs) and Coherent Matrix Multiplication Units (CMXUs). Taking cues from the proposed forward difference estimation of [<xref ref-type="bibr" rid="B72">72</xref>, <xref ref-type="bibr" rid="B127">127</xref>], FICONN demonstrates an advance in efficiency by perturbing all parameters at once in a random direction, rather than individually perturbing the parameters. The system implements a 3-layer DNN and achieves 92.7% test accuracy on vowel classification, comparable to digital computation results on similar tasks. FICONN&#x2019;s power consumption is dominated by thermal phase shifters, indicating that its performance can be improved even further with more efficient phase shifting units. Observations on the training curve and time to convergence indicate that this area presents a rich potential for comparison with training algorithms in standard hardware.</p>
</sec>
<sec id="s5-2">
<title>5.2 Further research</title>
<p>There is much room for research into alternative ways of maximizing the use of photonic components and building improved neural network designs around those novel abilities. Many approaches to date focus on replicating existing neural network architectures, but pushing the boundaries of photonic deep learning will require novel network designs that maximally exploit the strengths of photonics. One important avenue is to rethink the functions that are used within neural networks. Recent work has demonstrated the advantages of architectures that leverage spectral transforms applied globally to input data, particularly successful for PDE and scientific computing applications [<xref ref-type="bibr" rid="B195">195</xref>]. Implementing special functions and spectral transforms is more costly in digital hardware, which has so far restricted their utility for large-scale models. However, the inherent properties of photonic devices could make such models more computationally viable [<xref ref-type="bibr" rid="B39">39</xref>, <xref ref-type="bibr" rid="B196">196</xref>]. Further, the capacity of MZI meshes for encoding complex-valued operations opens the opportunity for applying complex-valued networks, which have important theoretical advantages but are currently impractical to implement in standard hardware [<xref ref-type="bibr" rid="B197">197</xref>].</p>
<p>Another direction is to push the boundaries of device configurations and component optimization through inverse design methods. Improved approaches can for instance enable more advanced design of reconfigurable and tunable components [<xref ref-type="bibr" rid="B198">198</xref>, <xref ref-type="bibr" rid="B199">199</xref>]. Recently, researchers have begun to explore the power of machine learning for inverse design. Deep learning shows promise for expanding possibilities in fast, robust, data-driven inverse design, opening the door for free-form approaches [<xref ref-type="bibr" rid="B200">200</xref>, <xref ref-type="bibr" rid="B201">201</xref>]. Applying deep learning-enhanced device design methods can, in turn, push the boundaries of what is possible in creating accelerators for deep learning.</p>
<p>Also, on the horizon of photonic accelerators is the field of quantum photonic ML accelerators (QPMLAs). Advances can occur both in the physical layer implementation using quantum memristive (QM) devices and in the improvement of algorithms that run on the application layer of the quantum fabric. A physical realization of a QM photonic system has been reported in [<xref ref-type="bibr" rid="B202">202</xref>]. The structure is based on classical photonic devices such as a tunable (dissipative) 50:50 beam splitter and an MZI assembled to realize quantum photonic characteristics without superconducting devices. This technique was adapted for integrated photonics by [<xref ref-type="bibr" rid="B203">203</xref>], realizing a reservoir computing model with three photons, nine modes, and three quantum memristors. This glass-based quantum machine was evaluated on both classical and quantum classification tasks, achieving over 95% accuracy, alongside additional capability for detecting quantum entanglement.</p>
<p>Also in the realm of Quantum Optical Neural Networks, Steinbrecher et al. [<xref ref-type="bibr" rid="B204">204</xref>] implement a design and observe how natural features of quantum optics can map to the operations of neural networks. They discuss how quantum optics can push beyond simply accelerating classical learning tasks by developing networks which are inherently modeled on quantum states such as coherence and entanglement, which is infeasible in classical computers. This can greatly enhance analysis on physical systems encoded by quantum information. Overall, integrating quantum capabilities into deep learning applications presents an exciting challenge for future developments in network design.</p>
</sec>
</sec>
<sec id="s6" sec-type="conclusion">
<title>6 Conclusion</title>
<p>Many deep learning operations can be greatly accelerated partially or entirely by photonic devices, allowing for remarkable speed and significantly lower energy consumption compared to their electronic counterparts. The increased compactness and integration density of state-of-the-art on-chip integrated photonics circuits have brought them into the realm of possibility for practical use in artificial neural networks. As shown by the examples in this study, photonic processors can be capable of performing deep learning inference with pre-trained networks at reduced power consumption and enhanced speed. Further, novel designs are even capable of training a model from scratch for end-to-end acceleration. OPUs still face persistent challenges in scalability and integration, and further progress is necessary in more holistic designs which fully integrate hardware, models, and algorithms. By promoting awareness in the deep learning community of the cutting-edge photonics capabilities, and knowledge of practical deep learning considerations in the photonics community, PDLA technology has the potential to circumvent existing resource constraints and expand the boundaries of AI applications.</p>
</sec>
</body>
<back>
<sec id="s7">
<title>Author contributions</title>
<p>MA: Conceptualization, Formal Analysis, Investigation, Project administration, Validation, Visualization, Writing&#x2013;original draft, Writing&#x2013;review and editing. SP: Writing&#x2013;review and editing, Conceptualization, Investigation, Visualization. SS: Validation, Visualization, Writing&#x2013;original draft, Writing&#x2013;review and editing. MR: Conceptualization, Formal Analysis, Project administration, Supervision, Writing&#x2013;original draft, Writing&#x2013;review and editing.</p>
</sec>
<sec id="s8" sec-type="funding-information">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research, authorship, and/or publication of this article. This work was made possible with the support of the NYUAD Research Enhancement Fund.</p>
</sec>
<ack>
<p>The authors express their gratitude to the NYU Abu Dhabi Center for Smart Engineering Materials and the Center for Cyber Security for their valuable contributions and support.</p>
</ack>
<sec id="s9" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s10" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<fn-group>
<fn id="fn3">
<label>1</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://lightmatter.co">https://lightmatter.co</ext-link>
</p>
</fn>
<fn id="fn4">
<label>2</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://www.lightelligence.ai">https://www.lightelligence.ai</ext-link>
</p>
</fn>
<fn id="fn5">
<label>3</label>
<p>
<ext-link ext-link-type="uri" xlink:href="https://www.luminous.com">https://www.luminous.com</ext-link>
</p>
</fn>
<fn id="fn6">
<label>4</label>
<p>Based on PyTorch library standard implementations, initialized with default weights pulled on 22 March 2024, code using the torchprofile utility on a forward pass of each network on a random tensor of size (32, 3, 224, 224) as a representative input batch size.</p>
</fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1">
<label>1.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schmidhuber</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Deep learning in neural networks: an overview</article-title>. <source>Neural networks</source> (<year>2015</year>) <volume>61</volume>:<fpage>85</fpage>&#x2013;<lpage>117</lpage>. <pub-id pub-id-type="doi">10.1016/j.neunet.2014.09.003</pub-id>
</citation>
</ref>
<ref id="B2">
<label>2.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Schmidhuber</surname>
<given-names>J</given-names>
</name>
</person-group>. <source>Annotated history of modern AI and deep learning</source> (<year>2022</year>). <pub-id pub-id-type="doi">10.48550/arXiv.2212.11279</pub-id>
</citation>
</ref>
<ref id="B3">
<label>3.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Moore</surname>
<given-names>GE</given-names>
</name>
</person-group>. <article-title>Progress in digital integrated electronics</article-title>. In: <source>International electron devices meeting (IEEE)</source> (<year>1975</year>). p. <fpage>11</fpage>&#x2013;<lpage>3</lpage>.</citation>
</ref>
<ref id="B4">
<label>4.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Krizhevsky</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Sutskever</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Hinton</surname>
<given-names>GE</given-names>
</name>
</person-group>. <article-title>Imagenet classification with deep convolutional neural networks</article-title>. <source>Adv Neural Inf Process Syst</source> (<year>2012</year>) <volume>25</volume>.</citation>
</ref>
<ref id="B5">
<label>5.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Chellapilla</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Puri</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Simard</surname>
<given-names>P</given-names>
</name>
</person-group>. <article-title>High performance convolutional neural networks for document processing</article-title>. In: <source>Tenth international workshop on frontiers in handwriting recognition (Suvisoft)</source> (<year>2006</year>).</citation>
</ref>
<ref id="B6">
<label>6.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cire&#x15f;an</surname>
<given-names>DC</given-names>
</name>
<name>
<surname>Meier</surname>
<given-names>U</given-names>
</name>
<name>
<surname>Masci</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Gambardella</surname>
<given-names>LM</given-names>
</name>
<name>
<surname>Schmidhuber</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Flexible, high performance convolutional neural networks for image classification</article-title>. <source>Proc Twenty-Second Int Jt Conf Artif Intelligence</source> (<year>2011</year>) <volume>2</volume>:<fpage>1237</fpage>&#x2013;<lpage>42</lpage>.</citation>
</ref>
<ref id="B7">
<label>7.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Fang</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Cong</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Caffeine: towards uniformed representation and acceleration for deep convolutional neural networks</article-title>. In: <source>2016 IEEE/ACM international conference on computer-aided design (ICCAD)</source>. <publisher-name>IEEE Press</publisher-name> (<year>2016</year>). p. <fpage>1</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1145/2966986.2967011</pub-id>
</citation>
</ref>
<ref id="B8">
<label>8.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Chiou</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>The microsoft catapult project</article-title>. In: <source>2017 IEEE international symposium on workload characterization (IISWC)</source>. <publisher-name>IEEE Computer Society</publisher-name> (<year>2017</year>). p. <fpage>124</fpage>.</citation>
</ref>
<ref id="B9">
<label>9.</label>
<citation citation-type="book">
<collab>NVIDIA A100 Tensor Core GPU Architecture</collab> <source>Whitepaper</source> (<year>2022</year>). <comment>NVIDIA (????)</comment>.</citation>
</ref>
<ref id="B10">
<label>10.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jordan</surname>
<given-names>K</given-names>
</name>
</person-group>. <article-title>94% on CIFAR-10 in 3.29 seconds on a single GPU</article-title> (<year>2024</year>). <pub-id pub-id-type="doi">10.48550/arXiv.2404.00498</pub-id>
</citation>
</ref>
<ref id="B11">
<label>11.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Cam</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Hungerford</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Schoch</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Pinto Miranda</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Y&#xe1;&#xf1;ez de Le&#xf3;n</surname>
<given-names>CD</given-names>
</name>
</person-group>. <source>Electricity 2024: analysis and forecast to 2026. Tech. rep</source>. <publisher-name>International Energy Agency</publisher-name> (<year>2024</year>).</citation>
</ref>
<ref id="B12">
<label>12.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>XY</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>XM</given-names>
</name>
</person-group>. <article-title>Integrated photonic computing beyond the von neumann architecture</article-title>. <source>ACS Photon</source> (<year>2023</year>) <volume>10</volume>:<fpage>1027</fpage>&#x2013;<lpage>36</lpage>. <pub-id pub-id-type="doi">10.1021/acsphotonics.2c01543</pub-id>
</citation>
</ref>
<ref id="B13">
<label>13.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rasras</surname>
<given-names>MS</given-names>
</name>
<name>
<surname>Gill</surname>
<given-names>DM</given-names>
</name>
<name>
<surname>Earnshaw</surname>
<given-names>MP</given-names>
</name>
<name>
<surname>Doerr</surname>
<given-names>CR</given-names>
</name>
<name>
<surname>Weiner</surname>
<given-names>JS</given-names>
</name>
<name>
<surname>Bolle</surname>
<given-names>CA</given-names>
</name>
<etal/>
</person-group> <article-title>Cmos silicon receiver integrated with ge detector and reconfigurable optical filter</article-title>. <source>IEEE Photon Tech Lett</source> (<year>2009</year>) <volume>22</volume>:<fpage>112</fpage>&#x2013;<lpage>4</lpage>. <pub-id pub-id-type="doi">10.1109/lpt.2009.2036590</pub-id>
</citation>
</ref>
<ref id="B14">
<label>14.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Melloni</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Martinelli</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>Synthesis of direct-coupled-resonators bandpass filters for wdm systems</article-title>. <source>J Lightwave Tech</source> (<year>2002</year>) <volume>20</volume>:<fpage>296</fpage>&#x2013;<lpage>303</lpage>. <pub-id pub-id-type="doi">10.1109/50.983244</pub-id>
</citation>
</ref>
<ref id="B15">
<label>15.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xiao</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Khan</surname>
<given-names>MH</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Qi</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>Multiple-channel silicon micro-resonator based filters for wdm applications</article-title>. <source>Opt Express</source> (<year>2007</year>) <volume>15</volume>:<fpage>7489</fpage>&#x2013;<lpage>98</lpage>. <pub-id pub-id-type="doi">10.1364/oe.15.007489</pub-id>
</citation>
</ref>
<ref id="B16">
<label>16.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheung</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Su</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Okamoto</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Yoo</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Ultra-compact silicon photonic 512 &#xd7; 512 25 ghz arrayed waveguide grating router</article-title>. <source>IEEE J Selected Top Quan Electron</source> (<year>2013</year>) <volume>20</volume>:<fpage>310</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1109/JSTQE.2013.2295879</pub-id>
</citation>
</ref>
<ref id="B17">
<label>17.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sorger</surname>
<given-names>VJ</given-names>
</name>
<name>
<surname>Lanzillotti-Kimura</surname>
<given-names>ND</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>RM</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X</given-names>
</name>
</person-group>. <article-title>Ultra-compact silicon nanophotonic modulator with broadband response</article-title>. <source>Nanophotonics</source> (<year>2012</year>) <volume>1</volume>:<fpage>17</fpage>&#x2013;<lpage>22</lpage>. <pub-id pub-id-type="doi">10.1515/nanoph-2012-0009</pub-id>
</citation>
</ref>
<ref id="B18">
<label>18.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Timurdogan</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Sorace-Agaskar</surname>
<given-names>CM</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Shah Hosseini</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Biberman</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Watts</surname>
<given-names>MR</given-names>
</name>
</person-group>. <article-title>An ultralow power athermal silicon modulator</article-title>. <source>Nat Commun</source> (<year>2014</year>) <volume>5</volume>:<fpage>4008</fpage>&#x2013;<lpage>11</lpage>. <pub-id pub-id-type="doi">10.1038/ncomms5008</pub-id>
</citation>
</ref>
<ref id="B19">
<label>19.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sepehrian</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Rusch</surname>
<given-names>LA</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>W</given-names>
</name>
</person-group>. <article-title>Silicon photonic iq modulators for 400 gb/s and beyond</article-title>. <source>J Lightwave Tech</source> (<year>2019</year>) <volume>37</volume>:<fpage>3078</fpage>&#x2013;<lpage>86</lpage>. <pub-id pub-id-type="doi">10.1109/jlt.2019.2910491</pub-id>
</citation>
</ref>
<ref id="B20">
<label>20.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rosenberg</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Green</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Assefa</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Gill</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Barwicz</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>M</given-names>
</name>
<etal/>
</person-group> <article-title>A 25 gbps silicon microring modulator based on an interleaved junction</article-title>. <source>Opt express</source> (<year>2012</year>) <volume>20</volume>:<fpage>26411</fpage>&#x2013;<lpage>23</lpage>. <pub-id pub-id-type="doi">10.1364/oe.20.026411</pub-id>
</citation>
</ref>
<ref id="B21">
<label>21.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ban</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Verbist</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Vanhoecke</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Bauwelinck</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Verheyen</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Lardenois</surname>
<given-names>S</given-names>
</name>
<etal/>
</person-group> <article-title>Low-voltage 60gb/s nrz and 100gb/s pam4 o-band silicon ring modulator</article-title>. In: <source>2019 IEEE optical interconnects conference (OI) (IEEE)</source> (<year>2019</year>). p. <fpage>1</fpage>&#x2013;<lpage>2</lpage>.</citation>
</ref>
<ref id="B22">
<label>22.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Javidi</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Tang</surname>
<given-names>Q</given-names>
</name>
</person-group>. <article-title>Optical implementation of neural networks for face recognition by the use of nonlinear joint transform correlators</article-title>. <source>Appl Opt</source> (<year>1995</year>) <volume>34</volume>:<fpage>3950</fpage>&#x2013;<lpage>62</lpage>. <pub-id pub-id-type="doi">10.1364/ao.34.003950</pub-id>
</citation>
</ref>
<ref id="B23">
<label>23.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Javidi</surname>
<given-names>B</given-names>
</name>
</person-group>. <article-title>Comparison of nonlinear joint transform correlator and nonlinearly transformed matched filter based correlator for noisy input scenes</article-title>. <source>Opt Eng</source> (<year>1990</year>) <volume>29</volume>:<fpage>1013</fpage>&#x2013;<lpage>20</lpage>. <pub-id pub-id-type="doi">10.1117/12.55703</pub-id>
</citation>
</ref>
<ref id="B24">
<label>24.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Reck</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Zeilinger</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Bernstein</surname>
<given-names>HJ</given-names>
</name>
<name>
<surname>Bertani</surname>
<given-names>P</given-names>
</name>
</person-group>. <article-title>Experimental realization of any discrete unitary operator</article-title>. <source>Phys Rev Lett</source> (<year>1994</year>) <volume>73</volume>:<fpage>58</fpage>&#x2013;<lpage>61</lpage>. <pub-id pub-id-type="doi">10.1103/physrevlett.73.58</pub-id>
</citation>
</ref>
<ref id="B25">
<label>25.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Miller</surname>
<given-names>DA</given-names>
</name>
</person-group>. <article-title>Establishing optimal wave communication channels automatically</article-title>. <source>J Lightwave Tech</source> (<year>2013</year>) <volume>31</volume>:<fpage>3987</fpage>&#x2013;<lpage>94</lpage>. <pub-id pub-id-type="doi">10.1109/jlt.2013.2278809</pub-id>
</citation>
</ref>
<ref id="B26">
<label>26.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Miller</surname>
<given-names>DA</given-names>
</name>
</person-group>. <article-title>Self-aligning universal beam coupler</article-title>. <source>Opt express</source> (<year>2013</year>) <volume>21</volume>:<fpage>6360</fpage>&#x2013;<lpage>70</lpage>. <pub-id pub-id-type="doi">10.1364/oe.21.006360</pub-id>
</citation>
</ref>
<ref id="B27">
<label>27.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Miller</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>Self-configuring universal linear optics</article-title>. <source>APS March Meet Abstr</source> (<year>2015</year>) <volume>2015</volume>:<fpage>S6</fpage>&#x2013;<lpage>001</lpage>.</citation>
</ref>
<ref id="B28">
<label>28.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Clements</surname>
<given-names>WR</given-names>
</name>
<name>
<surname>Humphreys</surname>
<given-names>PC</given-names>
</name>
<name>
<surname>Metcalf</surname>
<given-names>BJ</given-names>
</name>
<name>
<surname>Kolthammer</surname>
<given-names>WS</given-names>
</name>
<name>
<surname>Walmsley</surname>
<given-names>IA</given-names>
</name>
</person-group>. <article-title>Optimal design for universal multiport interferometers</article-title>. <source>Optica</source> (<year>2016</year>) <volume>3</volume>:<fpage>1460</fpage>&#x2013;<lpage>5</lpage>. <pub-id pub-id-type="doi">10.1364/OPTICA.3.001460</pub-id>
</citation>
</ref>
<ref id="B29">
<label>29.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Miller</surname>
<given-names>DA</given-names>
</name>
</person-group>. <article-title>Reconfigurable add-drop multiplexer for spatial modes</article-title>. <source>Opt express</source> (<year>2013</year>) <volume>21</volume>:<fpage>20220</fpage>&#x2013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.1364/oe.21.020220</pub-id>
</citation>
</ref>
<ref id="B30">
<label>30.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hardy</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Shamir</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Optics inspired logic architecture</article-title>. <source>Opt Express</source> (<year>2007</year>) <volume>15</volume>:<fpage>150</fpage>&#x2013;<lpage>65</lpage>. <pub-id pub-id-type="doi">10.1364/oe.15.000150</pub-id>
</citation>
</ref>
<ref id="B31">
<label>31.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schwelb</surname>
<given-names>O</given-names>
</name>
</person-group>. <article-title>Transmission, group delay, and dispersion in single-ring optical resonators and add/drop filters-a tutorial overview</article-title>. <source>J Lightwave Tech</source> (<year>2004</year>) <volume>22</volume>:<fpage>1380</fpage>&#x2013;<lpage>94</lpage>. <pub-id pub-id-type="doi">10.1109/jlt.2004.827666</pub-id>
</citation>
</ref>
<ref id="B32">
<label>32.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Shakya</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Lipson</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>Direct measurement of tunable optical delays on chip analogue to electromagnetically induced transparency</article-title>. <source>Opt express</source> (<year>2006</year>) <volume>14</volume>:<fpage>6463</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1364/oe.14.006463</pub-id>
</citation>
</ref>
<ref id="B33">
<label>33.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Fattal</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Beausoleil</surname>
<given-names>RG</given-names>
</name>
</person-group>. <article-title>Silicon microring resonators with 1.5-<italic>&#x3bc;</italic>m radius</article-title>. <source>Opt express</source> (<year>2008</year>) <volume>16</volume>:<fpage>4309</fpage>&#x2013;<lpage>15</lpage>. <pub-id pub-id-type="doi">10.1364/oe.16.004309</pub-id>
</citation>
</ref>
<ref id="B34">
<label>34.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Ji</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Jia</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>Y</given-names>
</name>
<etal/>
</person-group> <article-title>Demonstration of directed xor/xnor logic gates using two cascaded microring resonators</article-title>. <source>Opt Lett</source> (<year>2010</year>) <volume>35</volume>:<fpage>1620</fpage>&#x2013;<lpage>2</lpage>. <pub-id pub-id-type="doi">10.1364/ol.35.001620</pub-id>
</citation>
</ref>
<ref id="B35">
<label>35.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Takeuchi</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>Programmable phase-change metasurfaces on waveguides for multimode photonic convolutional neural network</article-title>. <source>Nat Commun</source> (<year>2021</year>) <volume>12</volume>:<fpage>96</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-020-20365-z</pub-id>
</citation>
</ref>
<ref id="B36">
<label>36.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tamalampudi</surname>
<given-names>SR</given-names>
</name>
<name>
<surname>Dushaq</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Villegas</surname>
<given-names>JE</given-names>
</name>
<name>
<surname>Rajput</surname>
<given-names>NS</given-names>
</name>
<name>
<surname>Paredes</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Elamurugu</surname>
<given-names>E</given-names>
</name>
<etal/>
</person-group> <article-title>Short-wavelength infrared (swir) photodetector based on multi-layer 2d gagete</article-title>. <source>Opt Express</source> (<year>2021</year>) <volume>29</volume>:<fpage>39395</fpage>&#x2013;<lpage>405</lpage>. <pub-id pub-id-type="doi">10.1364/oe.442845</pub-id>
</citation>
</ref>
<ref id="B37">
<label>37.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dushaq</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Paredes</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Villegas</surname>
<given-names>JE</given-names>
</name>
<name>
<surname>Tamalampudi</surname>
<given-names>SR</given-names>
</name>
<name>
<surname>Rasras</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>On-chip integration of 2d van der waals germanium phosphide (gep) for active silicon photonics devices</article-title>. <source>Opt Express</source> (<year>2022</year>) <volume>30</volume>:<fpage>15986</fpage>&#x2013;<lpage>97</lpage>. <pub-id pub-id-type="doi">10.1364/oe.457242</pub-id>
</citation>
</ref>
<ref id="B38">
<label>38.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tamalampudi</surname>
<given-names>SR</given-names>
</name>
<name>
<surname>Dushaq</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Villegas</surname>
<given-names>JE</given-names>
</name>
<name>
<surname>Paredes</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Rasras</surname>
<given-names>MS</given-names>
</name>
</person-group>. <article-title>A multi-layered gagete electro-optic device integrated in silicon photonics</article-title>. <source>J Lightwave Tech</source> (<year>2023</year>) <fpage>1</fpage>&#x2013;<lpage>7</lpage>. <pub-id pub-id-type="doi">10.1109/jlt.2023.3237818</pub-id>
</citation>
</ref>
<ref id="B39">
<label>39.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Serunjogi</surname>
<given-names>SM</given-names>
</name>
<name>
<surname>Sanduleanu</surname>
<given-names>MA</given-names>
</name>
<name>
<surname>Rasras</surname>
<given-names>MS</given-names>
</name>
</person-group>. <article-title>Volterra series based linearity analysis of a phase-modulated microwave photonic link</article-title>. <source>J Lightwave Tech</source> (<year>2017</year>) <volume>36</volume>:<fpage>1537</fpage>&#x2013;<lpage>51</lpage>. <pub-id pub-id-type="doi">10.1109/JLT.2017.2782886</pub-id>
</citation>
</ref>
<ref id="B40">
<label>40.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Psaltis</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Brady</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Wagner</surname>
<given-names>K</given-names>
</name>
</person-group>. <article-title>Adaptive optical networks using photorefractive crystals</article-title>. <source>Appl Opt</source> (<year>1988</year>) <volume>27</volume>:<fpage>1752</fpage>&#x2013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.1364/ao.27.001752</pub-id>
</citation>
</ref>
<ref id="B41">
<label>41.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Farhat</surname>
<given-names>NH</given-names>
</name>
<name>
<surname>Psaltis</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Prata</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Paek</surname>
<given-names>E</given-names>
</name>
</person-group>. <article-title>Optical implementation of the Hopfield model</article-title>. <source>Appl Opt</source> (<year>1985</year>) <volume>24</volume>:<fpage>1469</fpage>. <pub-id pub-id-type="doi">10.1364/AO.24.001469</pub-id>
</citation>
</ref>
<ref id="B42">
<label>42.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ito</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Ki</surname>
<given-names>K</given-names>
</name>
</person-group>. <article-title>Optical implementation of the Hopfield neural network using multiple fiber nets</article-title>. <source>Appl Opt</source> (<year>1989</year>) <volume>28</volume>:<fpage>4176</fpage>. <pub-id pub-id-type="doi">10.1364/AO.28.004176</pub-id>
</citation>
</ref>
<ref id="B43">
<label>43.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Choquette</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Gandhi</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Giroux</surname>
<given-names>O</given-names>
</name>
<name>
<surname>Stam</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Krashinsky</surname>
<given-names>R</given-names>
</name>
</person-group>. <article-title>Nvidia a100 tensor core gpu: performance and innovation</article-title>. <source>IEEE Micro</source> (<year>2021</year>) <volume>41</volume>:<fpage>29</fpage>&#x2013;<lpage>35</lpage>. <pub-id pub-id-type="doi">10.1109/mm.2021.3061394</pub-id>
</citation>
</ref>
<ref id="B44">
<label>44.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>James</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>Iedm 2017: intel&#x2019;s 10nm platform process</article-title>. In: <source>Solid state technology</source> (<year>2017</year>).</citation>
</ref>
<ref id="B45">
<label>45.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Ye</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Lou</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>L</given-names>
</name>
</person-group>. <article-title>Holylight: a nanophotonic accelerator for deep learning in data centers</article-title>. In: <source>2019 design, automation &#x26; test in europe conference &#x26; exhibition (DATE)</source> (<year>2019</year>). p. <fpage>1483</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.23919/DATE.2019.8715195</pub-id>
</citation>
</ref>
<ref id="B46">
<label>46.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Fujiwara</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Mori</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>WC</given-names>
</name>
<name>
<surname>Chuang</surname>
<given-names>MC</given-names>
</name>
<name>
<surname>Naous</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Chuang</surname>
<given-names>CK</given-names>
</name>
<etal/>
</person-group> <article-title>A 5-nm 254-tops/w 221-tops/mm 2 fully-digital computing-in-memory macro supporting wide-range dynamic-voltage-frequency scaling and simultaneous mac and write operations</article-title>. In: <source>2022 IEEE international solid-state circuits conference (ISSCC) (IEEE)</source>, <volume>65</volume> (<year>2022</year>). p. <fpage>1</fpage>&#x2013;<lpage>3</lpage>. <pub-id pub-id-type="doi">10.1109/isscc42614.2022.9731754</pub-id>
</citation>
</ref>
<ref id="B47">
<label>47.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Mori</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>WC</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>CE</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>CF</given-names>
</name>
<name>
<surname>Hsu</surname>
<given-names>YH</given-names>
</name>
<name>
<surname>Chuang</surname>
<given-names>CK</given-names>
</name>
<etal/>
</person-group> <article-title>A 4nm 6163-tops/w/b <bold>4790 &#x2212; TOPS/mm</bold>
<sup>2</sup>
<bold>/b</bold> sram based digital-computing-in-memory macro supporting bit-width flexibility and simultaneous mac and weight update</article-title>. In: <source>2023 IEEE international solid-state circuits conference (ISSCC) (IEEE)</source> (<year>2023</year>). p. <fpage>132</fpage>&#x2013;<lpage>4</lpage>.</citation>
</ref>
<ref id="B48">
<label>48.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Farmakidis</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Youngblood</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Tan</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Swett</surname>
<given-names>JL</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>Z</given-names>
</name>
<etal/>
</person-group> <article-title>Plasmonic nanogap enhanced phase-change devices with dual electrical-optical functionality</article-title>. <source>Sci Adv</source> (<year>2019</year>) <volume>5</volume>:<fpage>eaaw2687</fpage>. <pub-id pub-id-type="doi">10.1126/sciadv.aaw2687</pub-id>
</citation>
</ref>
<ref id="B49">
<label>49.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Lu</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>H</given-names>
</name>
<etal/>
</person-group> <article-title>Miniature multilevel optical memristive switch using phase change material</article-title>. <source>ACS Photon</source> (<year>2019</year>) <volume>6</volume>:<fpage>2205</fpage>&#x2013;<lpage>12</lpage>. <pub-id pub-id-type="doi">10.1021/acsphotonics.9b00819</pub-id>
</citation>
</ref>
<ref id="B50">
<label>50.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Feldmann</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Youngblood</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Wright</surname>
<given-names>CD</given-names>
</name>
<name>
<surname>Bhaskaran</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Pernice</surname>
<given-names>WH</given-names>
</name>
</person-group>. <article-title>Integrated 256 cell photonic phase-change memory with 512-bit capacity</article-title>. <source>IEEE J Selected Top Quan Electron</source> (<year>2019</year>) <volume>26</volume>:<fpage>1</fpage>&#x2013;<lpage>7</lpage>. <pub-id pub-id-type="doi">10.1109/jstqe.2019.2956871</pub-id>
</citation>
</ref>
<ref id="B51">
<label>51.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tait</surname>
<given-names>AN</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>AX</given-names>
</name>
<name>
<surname>De Lima</surname>
<given-names>TF</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Shastri</surname>
<given-names>BJ</given-names>
</name>
<name>
<surname>Nahmias</surname>
<given-names>MA</given-names>
</name>
<etal/>
</person-group> <article-title>Microring weight banks</article-title>. <source>IEEE J Selected Top Quan Electron</source> (<year>2016</year>) <volume>22</volume>:<fpage>312</fpage>&#x2013;<lpage>25</lpage>. <pub-id pub-id-type="doi">10.1109/jstqe.2016.2573583</pub-id>
</citation>
</ref>
<ref id="B52">
<label>52.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Farmakidis</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Feldmann</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Tan</surname>
<given-names>J</given-names>
</name>
<name>
<surname>He</surname>
<given-names>Y</given-names>
</name>
<etal/>
</person-group> <article-title>Phase-change materials for energy-efficient photonic memory and computing</article-title>. <source>MRS Bull</source> (<year>2022</year>) <volume>47</volume>:<fpage>502</fpage>&#x2013;<lpage>10</lpage>. <pub-id pub-id-type="doi">10.1557/s43577-022-00358-7</pub-id>
</citation>
</ref>
<ref id="B53">
<label>53.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Miscuglio</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Adam</surname>
<given-names>GC</given-names>
</name>
<name>
<surname>Kuzum</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Sorger</surname>
<given-names>VJ</given-names>
</name>
</person-group>. <article-title>Roadmap on material-function mapping for photonic-electronic hybrid neural networks</article-title>. <source>APL Mater</source> (<year>2019</year>) <volume>7</volume>. <pub-id pub-id-type="doi">10.1063/1.5109689</pub-id>
</citation>
</ref>
<ref id="B54">
<label>54.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ma</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Meng</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Peserico</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Miscuglio</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>J</given-names>
</name>
<etal/>
</person-group> <article-title>Photonic tensor core with photonic compute-in-memory</article-title>. In: <source>Optical fiber communication conference</source>. <publisher-name>Optica Publishing Group</publisher-name> (<year>2022</year>). p. <fpage>M2E</fpage>&#x2013;<lpage>4</lpage>.</citation>
</ref>
<ref id="B55">
<label>55.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Peserico</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Shastri</surname>
<given-names>BJ</given-names>
</name>
<name>
<surname>Sorger</surname>
<given-names>VJ</given-names>
</name>
</person-group>. <article-title>Photonic tensor core for machine learning: a review</article-title>. In: <source>Emerging topics in artificial intelligence (ETAI) 2022 12204</source> (<year>2022</year>). p. <fpage>53</fpage>&#x2013;<lpage>60</lpage>.</citation>
</ref>
<ref id="B56">
<label>56.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>R&#xed;os</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Youngblood</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Le Gallo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Pernice</surname>
<given-names>WH</given-names>
</name>
<name>
<surname>Wright</surname>
<given-names>CD</given-names>
</name>
<etal/>
</person-group> <article-title>In-memory computing on a photonic platform</article-title>. <source>Sci Adv</source> (<year>2019</year>) <volume>5</volume>:<fpage>eaau5759</fpage>. <pub-id pub-id-type="doi">10.1126/sciadv.aau5759</pub-id>
</citation>
</ref>
<ref id="B57">
<label>57.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Takeuchi</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>Programmable phase-change metasurface for multimode photonic convolutional neural network</article-title>. In: <source>2020 IEEE photonics conference (IPC) (IEEE)</source> (<year>2020</year>). p. <fpage>1</fpage>&#x2013;<lpage>2</lpage>.</citation>
</ref>
<ref id="B58">
<label>58.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>R&#xed;os</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Youngblood</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Wright</surname>
<given-names>CD</given-names>
</name>
<name>
<surname>Pernice</surname>
<given-names>WH</given-names>
</name>
<name>
<surname>Bhaskaran</surname>
<given-names>H</given-names>
</name>
</person-group>. <article-title>Device-level photonic memories and logic applications using phase-change materials</article-title>. <source>Adv Mater</source> (<year>2018</year>) <volume>30</volume>:<fpage>1802435</fpage>. <pub-id pub-id-type="doi">10.1002/adma.201802435</pub-id>
</citation>
</ref>
<ref id="B59">
<label>59.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Chou</surname>
<given-names>JB</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Du</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Yadav</surname>
<given-names>A</given-names>
</name>
<etal/>
</person-group> <article-title>Broadband transparent optical phase change materials for high-performance nonvolatile photonics</article-title>. <source>Nat Commun</source> (<year>2019</year>) <volume>10</volume>:<fpage>4279</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-019-12196-4</pub-id>
</citation>
</ref>
<ref id="B60">
<label>60.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sebastian</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Le Gallo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Burr</surname>
<given-names>GW</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>S</given-names>
</name>
<name>
<surname>BrightSky</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Eleftheriou</surname>
<given-names>E</given-names>
</name>
</person-group>. <article-title>Tutorial: brain-inspired computing using phase-change memory devices</article-title>. <source>J Appl Phys</source> (<year>2018</year>) <volume>124</volume>. <pub-id pub-id-type="doi">10.1063/1.5042413</pub-id>
</citation>
</ref>
<ref id="B61">
<label>61.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ambrogio</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Narayanan</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Tsai</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Shelby</surname>
<given-names>RM</given-names>
</name>
<name>
<surname>Boybat</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Di</surname>
<given-names>NC</given-names>
</name>
<etal/>
</person-group> <article-title>Equivalent-accuracy accelerated neural-network training using analogue memory</article-title>. <source>Nature</source> (<year>2018</year>) <volume>558</volume>:<fpage>60</fpage>&#x2013;<lpage>7</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-018-0180-5</pub-id>
</citation>
</ref>
<ref id="B62">
<label>62.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Farquhar</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Hasler</surname>
<given-names>P</given-names>
</name>
</person-group>. <article-title>A bio-physically inspired silicon neuron</article-title>. <source>IEEE Trans Circuits Syst Regular Pap</source> (<year>2005</year>) <volume>52</volume>:<fpage>477</fpage>&#x2013;<lpage>88</lpage>. <pub-id pub-id-type="doi">10.1109/tcsi.2004.842871</pub-id>
</citation>
</ref>
<ref id="B63">
<label>63.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Szilagyi</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Pliva</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Henker</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Schoeniger</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Turkiewicz</surname>
<given-names>JP</given-names>
</name>
<name>
<surname>Ellinger</surname>
<given-names>F</given-names>
</name>
</person-group>. <article-title>A 53-gbit/s optical receiver frontend with 0.65 pj/bit in 28-nm bulk-cmos</article-title>. <source>IEEE J Solid-State Circuits</source> (<year>2018</year>) <volume>54</volume>:<fpage>845</fpage>&#x2013;<lpage>55</lpage>. <pub-id pub-id-type="doi">10.1109/jssc.2018.2885531</pub-id>
</citation>
</ref>
<ref id="B64">
<label>64.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Stroev</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Berloff</surname>
<given-names>NG</given-names>
</name>
</person-group>. <article-title>Analog photonics computing for information processing, inference, and optimization</article-title>. <source>Adv Quan Tech</source> (<year>2023</year>) <volume>6</volume>:<fpage>2300055</fpage>. <pub-id pub-id-type="doi">10.1002/qute.202300055</pub-id>
</citation>
</ref>
<ref id="B65">
<label>65.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Sorger</surname>
<given-names>VJ</given-names>
</name>
<name>
<surname>Miscuglio</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Al-Qadasi</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Mukherjee</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Lampe</surname>
<given-names>L</given-names>
</name>
<etal/>
</person-group> <article-title>Prospects and applications of photonic neural networks</article-title>. <source>Adv Phys X</source> (<year>2022</year>) <volume>7</volume>:<fpage>1981155</fpage>. <pub-id pub-id-type="doi">10.1080/23746149.2021.1981155</pub-id>
</citation>
</ref>
<ref id="B66">
<label>66.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Fang</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Jiao</surname>
<given-names>S</given-names>
</name>
<etal/>
</person-group> <article-title>Analog optical computing for artificial intelligence</article-title>. <source>Engineering</source> (<year>2022</year>) <volume>10</volume>:<fpage>133</fpage>&#x2013;<lpage>45</lpage>. <pub-id pub-id-type="doi">10.1016/j.eng.2021.06.021</pub-id>
</citation>
</ref>
<ref id="B67">
<label>67.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Al-Qadasi</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Chrostowski</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Shastri</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Shekhar</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Scaling up silicon photonic-based accelerators: challenges and opportunities</article-title>. <source>APL Photon</source> (<year>2022</year>) <volume>7</volume>. <pub-id pub-id-type="doi">10.1063/5.0070992</pub-id>
</citation>
</ref>
<ref id="B68">
<label>68.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Xia</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Photonic computing and communication for neural network accelerators</article-title>. In: <source>International conference on parallel and distributed computing: applications and technologies</source>. <publisher-name>Springer</publisher-name> (<year>2021</year>). p. <fpage>121</fpage>&#x2013;<lpage>8</lpage>.</citation>
</ref>
<ref id="B69">
<label>69.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ma</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Peserico</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Khaled</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Nouri</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Dalir</surname>
<given-names>H</given-names>
</name>
<etal/>
</person-group> <source>High-density integrated photonic tensor processing unit with a matrix multiply compiler</source> (<year>2022</year>).</citation>
</ref>
<ref id="B70">
<label>70.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Launay</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Poli</surname>
<given-names>I</given-names>
</name>
<name>
<surname>M&#xfc;ller</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Carron</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Daudet</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Krzakala</surname>
<given-names>F</given-names>
</name>
<etal/>
</person-group> <source>Light-in-the-loop: using a photonics co-processor for scalable training of neural networks</source> (<year>2020</year>). <comment>arXiv preprint arXiv:2006.01475</comment>.</citation>
</ref>
<ref id="B71">
<label>71.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Hesslow</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Cappelli</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Carron</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Daudet</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Lafargue</surname>
<given-names>R</given-names>
</name>
<name>
<surname>M&#xfc;ller</surname>
<given-names>K</given-names>
</name>
<etal/>
</person-group> <source>Photonic co-processors in hpc: using lighton opus for randomized numerical linear algebra</source> (<year>2021</year>). <comment>arXiv preprint arXiv:2104.14429</comment>.</citation>
</ref>
<ref id="B72">
<label>72.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hughes</surname>
<given-names>TW</given-names>
</name>
<name>
<surname>Minkov</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Fan</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Training of photonic neural networks through <italic>in situ</italic> backpropagation and gradient measurement</article-title>. <source>Optica</source> (<year>2018</year>) <volume>5</volume>:<fpage>864</fpage>&#x2013;<lpage>71</lpage>. <pub-id pub-id-type="doi">10.1364/optica.5.000864</pub-id>
</citation>
</ref>
<ref id="B73">
<label>73.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Brossollet</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Cappelli</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Carron</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Chaintoutis</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Chatelain</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Daudet</surname>
<given-names>L</given-names>
</name>
<etal/>
</person-group> <source>Lighton optical processing unit: Scaling-up ai and hpc with a non von neumann co-processor</source> (<year>2021</year>). <comment>arXiv preprint arXiv:2107.11814</comment>.</citation>
</ref>
<ref id="B74">
<label>74.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lu</surname>
<given-names>L</given-names>
</name>
</person-group>. <article-title>Elighting up the future</article-title>. <source>Light: Sci Appl</source> (<year>2021</year>) <volume>10</volume>:<fpage>118</fpage>. <pub-id pub-id-type="doi">10.1038/s41377-021-00555-0</pub-id>
</citation>
</ref>
<ref id="B75">
<label>75.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Burr</surname>
<given-names>GW</given-names>
</name>
<name>
<surname>Brightsky</surname>
<given-names>MJ</given-names>
</name>
<name>
<surname>Sebastian</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>HY</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>JY</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>S</given-names>
</name>
<etal/>
</person-group> <article-title>Recent progress in phase-change memory technology</article-title>. <source>IEEE J Emerging Selected Top Circuits Syst</source> (<year>2016</year>) <volume>6</volume>:<fpage>146</fpage>&#x2013;<lpage>62</lpage>.</citation>
</ref>
<ref id="B76">
<label>76.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Harris</surname>
<given-names>NC</given-names>
</name>
<name>
<surname>Carolan</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Bunandar</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Prabhu</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Hochberg</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Baehr-Jones</surname>
<given-names>T</given-names>
</name>
<etal/>
</person-group> <article-title>Linear programmable nanophotonic processors</article-title>. <source>Optica</source> (<year>2018</year>) <volume>5</volume>:<fpage>1623</fpage>&#x2013;<lpage>31</lpage>. <pub-id pub-id-type="doi">10.1364/OPTICA.5.001623</pub-id>
</citation>
</ref>
<ref id="B77">
<label>77.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bogaerts</surname>
<given-names>W</given-names>
</name>
<name>
<surname>P&#xe9;rez</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Capmany</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Miller</surname>
<given-names>DA</given-names>
</name>
<name>
<surname>Poon</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Englund</surname>
<given-names>D</given-names>
</name>
<etal/>
</person-group> <article-title>Programmable photonic circuits</article-title>. <source>Nature</source> (<year>2020</year>) <volume>586</volume>:<fpage>207</fpage>&#x2013;<lpage>16</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-020-2764-0</pub-id>
</citation>
</ref>
<ref id="B78">
<label>78.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Youngblood</surname>
<given-names>N</given-names>
</name>
<name>
<surname>R&#xed;os</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Wright</surname>
<given-names>CD</given-names>
</name>
<name>
<surname>Pernice</surname>
<given-names>WH</given-names>
</name>
<etal/>
</person-group> <article-title>Fast and reliable storage using a 5 bit, nonvolatile photonic memory cell</article-title>. <source>Optica</source> (<year>2019</year>) <volume>6</volume>:<fpage>1</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1364/optica.6.000001</pub-id>
</citation>
</ref>
<ref id="B79">
<label>79.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Zheng</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Doylend</surname>
<given-names>JK</given-names>
</name>
<name>
<surname>Majumdar</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Low-loss and broadband nonvolatile phase-change directional coupler switches</article-title>. <source>Acs Photon</source> (<year>2019</year>) <volume>6</volume>:<fpage>553</fpage>&#x2013;<lpage>7</lpage>. <pub-id pub-id-type="doi">10.1021/acsphotonics.8b01628</pub-id>
</citation>
</ref>
<ref id="B80">
<label>80.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wuttig</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Bhaskaran</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Taubner</surname>
<given-names>T</given-names>
</name>
</person-group>. <article-title>Phase-change materials for non-volatile photonic applications</article-title>. <source>Nat Photon</source> (<year>2017</year>) <volume>11</volume>:<fpage>465</fpage>&#x2013;<lpage>76</lpage>. <pub-id pub-id-type="doi">10.1038/nphoton.2017.126</pub-id>
</citation>
</ref>
<ref id="B81">
<label>81.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yang</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Ramanathan</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Breakthroughs in photonics 2014: phase change materials for photonics</article-title>. <source>IEEE Photon J</source> (<year>2015</year>) <volume>7</volume>:<fpage>1</fpage>&#x2013;<lpage>5</lpage>. <pub-id pub-id-type="doi">10.1109/jphot.2015.2413594</pub-id>
</citation>
</ref>
<ref id="B82">
<label>82.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Soref</surname>
<given-names>R</given-names>
</name>
</person-group>. <article-title>Reconfigurable optical directed-logic circuits using microresonator-based optical switches</article-title>. <source>Opt Express</source> (<year>2011</year>) <volume>19</volume>:<fpage>5244</fpage>&#x2013;<lpage>59</lpage>. <pub-id pub-id-type="doi">10.1364/oe.19.005244</pub-id>
</citation>
</ref>
<ref id="B83">
<label>83.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Luo</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Wan</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>S</given-names>
</name>
<etal/>
</person-group> <article-title>Recent progress in quantum photonic chips for quantum communication and internet</article-title>. <source>Light: Sci Appl</source> (<year>2023</year>) <volume>12</volume>:<fpage>175</fpage>. <pub-id pub-id-type="doi">10.1038/s41377-023-01173-8</pub-id>
</citation>
</ref>
<ref id="B84">
<label>84.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Paraiso</surname>
<given-names>TK</given-names>
</name>
<name>
<surname>Roger</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Marangon</surname>
<given-names>DG</given-names>
</name>
<name>
<surname>De Marco</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Sanzaro</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Woodward</surname>
<given-names>RI</given-names>
</name>
<etal/>
</person-group> <article-title>A photonic integrated quantum secure communication system</article-title>. <source>Nat Photon</source> (<year>2021</year>) <volume>15</volume>:<fpage>850</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1038/s41566-021-00873-0</pub-id>
</citation>
</ref>
<ref id="B85">
<label>85.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Litvin</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Martynenko</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Purcell-Milton</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Baranov</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Fedorov</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Gun&#x2019;Ko</surname>
<given-names>Y</given-names>
</name>
</person-group>. <article-title>Colloidal quantum dots for optoelectronics</article-title>. <source>J Mater Chem A</source> (<year>2017</year>) <volume>5</volume>:<fpage>13252</fpage>&#x2013;<lpage>75</lpage>. <pub-id pub-id-type="doi">10.1039/c7ta02076g</pub-id>
</citation>
</ref>
<ref id="B86">
<label>86.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Tate</surname>
<given-names>N</given-names>
</name>
</person-group>. <article-title>Quantum-dot-based photonic reservoir computing</article-title>. In: <source>Photonic neural networks with spatiotemporal dynamics</source> (<year>2024</year>). p. <fpage>71</fpage>.</citation>
</ref>
<ref id="B87">
<label>87.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lingnau</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Perrott</surname>
<given-names>AH</given-names>
</name>
<name>
<surname>Dernaika</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Caro</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Peters</surname>
<given-names>FH</given-names>
</name>
<name>
<surname>Kelleher</surname>
<given-names>B</given-names>
</name>
</person-group>. <article-title>Dynamics of on-chip asymmetrically coupled semiconductor lasers</article-title>. <source>Opt Lett</source> (<year>2020</year>) <volume>45</volume>:<fpage>2223</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1364/ol.390401</pub-id>
</citation>
</ref>
<ref id="B88">
<label>88.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shainline</surname>
<given-names>JM</given-names>
</name>
<name>
<surname>Buckley</surname>
<given-names>SM</given-names>
</name>
<name>
<surname>Mirin</surname>
<given-names>RP</given-names>
</name>
<name>
<surname>Nam</surname>
<given-names>SW</given-names>
</name>
</person-group>. <article-title>Superconducting optoelectronic circuits for neuromorphic computing</article-title>. <source>Phys Rev Appl</source> (<year>2017</year>) <volume>7</volume>:<fpage>034013</fpage>. <pub-id pub-id-type="doi">10.1103/physrevapplied.7.034013</pub-id>
</citation>
</ref>
<ref id="B89">
<label>89.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Paesani</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Ding</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Santagati</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Skrzypczyk</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Salavrakos</surname>
<given-names>A</given-names>
</name>
<etal/>
</person-group> <article-title>Multidimensional quantum entanglement with large-scale integrated optics</article-title>. <source>Science</source> (<year>2018</year>) <volume>360</volume>:<fpage>285</fpage>&#x2013;<lpage>91</lpage>. <pub-id pub-id-type="doi">10.1126/science.aar7053</pub-id>
</citation>
</ref>
<ref id="B90">
<label>90.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Politi</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Matthews</surname>
<given-names>JC</given-names>
</name>
<name>
<surname>O&#x2019;Brien</surname>
<given-names>JL</given-names>
</name>
</person-group>. <article-title>Shor&#x2019;s quantum factoring algorithm on a photonic chip</article-title>. <source>Science</source> (<year>2009</year>) <volume>325</volume>:<fpage>1221</fpage>. <pub-id pub-id-type="doi">10.1126/science.1173731</pub-id>
</citation>
</ref>
<ref id="B91">
<label>91.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>XQ</given-names>
</name>
<name>
<surname>Kalasuwan</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Ralph</surname>
<given-names>TC</given-names>
</name>
<name>
<surname>O&#x2019;brien</surname>
<given-names>JL</given-names>
</name>
</person-group>. <article-title>Calculating unknown eigenvalues with a quantum algorithm</article-title>. <source>Nat Photon</source> (<year>2013</year>) <volume>7</volume>:<fpage>223</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1038/nphoton.2012.360</pub-id>
</citation>
</ref>
<ref id="B92">
<label>92.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qiang</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Xue</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Ge</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y</given-names>
</name>
<etal/>
</person-group> <article-title>Implementing graph-theoretic quantum algorithms on a silicon photonic quantum walk processor</article-title>. <source>Sci Adv</source> (<year>2021</year>) <volume>7</volume>:<fpage>eabb8375</fpage>. <pub-id pub-id-type="doi">10.1126/sciadv.abb8375</pub-id>
</citation>
</ref>
<ref id="B93">
<label>93.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Sciarrino</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Laing</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Thompson</surname>
<given-names>MG</given-names>
</name>
</person-group>. <article-title>Integrated photonic quantum technologies</article-title>. <source>Nat Photon</source> (<year>2020</year>) <volume>14</volume>:<fpage>273</fpage>&#x2013;<lpage>84</lpage>. <pub-id pub-id-type="doi">10.1038/s41566-019-0532-1</pub-id>
</citation>
</ref>
<ref id="B94">
<label>94.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hsu</surname>
<given-names>CY</given-names>
</name>
<name>
<surname>Yiu</surname>
<given-names>GZ</given-names>
</name>
<name>
<surname>Chang</surname>
<given-names>YC</given-names>
</name>
</person-group>. <article-title>Free-space applications of silicon photonics: a review</article-title>. <source>Micromachines</source> (<year>2022</year>) <volume>13</volume>:<fpage>990</fpage>. <pub-id pub-id-type="doi">10.3390/mi13070990</pub-id>
</citation>
</ref>
<ref id="B95">
<label>95.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhu</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Lu</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>L</given-names>
</name>
</person-group>. <article-title>Design and experimental verification for optical module of optical vector&#x2013;matrix multiplier</article-title>. <source>Appl Opt</source> (<year>2013</year>) <volume>52</volume>:<fpage>4412</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1364/ao.52.004412</pub-id>
</citation>
</ref>
<ref id="B96">
<label>96.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fontaine</surname>
<given-names>NK</given-names>
</name>
<name>
<surname>Ryf</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Neilson</surname>
<given-names>DT</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Carpenter</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Laguerre-Gaussian mode sorter</article-title>. <source>Nat Commun</source> (<year>2019</year>) <volume>10</volume>:<fpage>1865</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-019-09840-4</pub-id>
</citation>
</ref>
<ref id="B97">
<label>97.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lin</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Rivenson</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Yardimci</surname>
<given-names>NT</given-names>
</name>
<name>
<surname>Veli</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Luo</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Jarrahi</surname>
<given-names>M</given-names>
</name>
<etal/>
</person-group> <article-title>All-optical machine learning using diffractive deep neural networks</article-title>. <source>Science</source> (<year>2018</year>) <volume>361</volume>:<fpage>1004</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1126/science.aat8084</pub-id>
</citation>
</ref>
<ref id="B98">
<label>98.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hamerly</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Bernstein</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Sludds</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Solja&#x10d;i&#x107;</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Englund</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>Large-scale optical neural networks based on photoelectric multiplication</article-title>. <source>Phys Rev X</source> (<year>2019</year>) <volume>9</volume>:<fpage>021032</fpage>. <pub-id pub-id-type="doi">10.1103/physrevx.9.021032</pub-id>
</citation>
</ref>
<ref id="B99">
<label>99.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>Y</given-names>
</name>
<etal/>
</person-group> <article-title>Photonic matrix multiplication lights up photonic accelerator and beyond</article-title>. <source>Light: Sci Appl</source> (<year>2022</year>) <volume>11</volume>:<fpage>30</fpage>. <pub-id pub-id-type="doi">10.1038/s41377-022-00717-8</pub-id>
</citation>
</ref>
<ref id="B100">
<label>100.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Y</given-names>
</name>
<etal/>
</person-group> <article-title>Large-scale neuromorphic optoelectronic computing with a reconfigurable diffractive processing unit</article-title>. <source>Nat Photon</source> (<year>2021</year>) <volume>15</volume>:<fpage>367</fpage>&#x2013;<lpage>73</lpage>. <pub-id pub-id-type="doi">10.1038/s41566-021-00796-w</pub-id>
</citation>
</ref>
<ref id="B101">
<label>101.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cordaro</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Edwards</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Nikkhah</surname>
<given-names>V</given-names>
</name>
<name>
<surname>Al&#xf9;</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Engheta</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Polman</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Solving integral equations in free space with inverse-designed ultrathin optical metagratings</article-title>. <source>Nat Nanotechnology</source> (<year>2023</year>) <volume>18</volume>:<fpage>365</fpage>&#x2013;<lpage>72</lpage>. <pub-id pub-id-type="doi">10.1038/s41565-022-01297-9</pub-id>
</citation>
</ref>
<ref id="B102">
<label>102.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Deng</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Tang</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Aza&#xf1;a</surname>
<given-names>J</given-names>
</name>
<etal/>
</person-group> <article-title>Reconfigurable optical signal processing based on a distributed feedback semiconductor optical amplifier</article-title>. <source>Scientific Rep</source> (<year>2016</year>) <volume>6</volume>:<fpage>19985</fpage>. <pub-id pub-id-type="doi">10.1038/srep19985</pub-id>
</citation>
</ref>
<ref id="B103">
<label>103.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tanaka</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Yamane</surname>
<given-names>T</given-names>
</name>
<name>
<surname>H&#xe9;roux</surname>
<given-names>JB</given-names>
</name>
<name>
<surname>Nakane</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Kanazawa</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Takeda</surname>
<given-names>S</given-names>
</name>
<etal/>
</person-group> <article-title>Recent advances in physical reservoir computing: a review</article-title>. <source>Neural Networks</source> (<year>2019</year>) <volume>115</volume>:<fpage>100</fpage>&#x2013;<lpage>23</lpage>. <pub-id pub-id-type="doi">10.1016/j.neunet.2019.03.005</pub-id>
</citation>
</ref>
<ref id="B104">
<label>104.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Paquot</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Duport</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Smerieri</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Dambre</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Schrauwen</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Haelterman</surname>
<given-names>M</given-names>
</name>
<etal/>
</person-group> <article-title>Optoelectronic reservoir computing</article-title>. <source>Scientific Rep</source> (<year>2012</year>) <volume>2</volume>:<fpage>287</fpage>. <pub-id pub-id-type="doi">10.1038/srep00287</pub-id>
</citation>
</ref>
<ref id="B105">
<label>105.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vandoorne</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Dambre</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Verstraeten</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Schrauwen</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Bienstman</surname>
<given-names>P</given-names>
</name>
</person-group>. <article-title>Parallel reservoir computing using optical amplifiers</article-title>. <source>IEEE Trans Neural networks</source> (<year>2011</year>) <volume>22</volume>:<fpage>1469</fpage>&#x2013;<lpage>81</lpage>. <pub-id pub-id-type="doi">10.1109/tnn.2011.2161771</pub-id>
</citation>
</ref>
<ref id="B106">
<label>106.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Larger</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Soriano</surname>
<given-names>MC</given-names>
</name>
<name>
<surname>Brunner</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Appeltant</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Guti&#xe9;rrez</surname>
<given-names>JM</given-names>
</name>
<name>
<surname>Pesquera</surname>
<given-names>L</given-names>
</name>
<etal/>
</person-group> <article-title>Photonic information processing beyond turing: an optoelectronic implementation of reservoir computing</article-title>. <source>Opt express</source> (<year>2012</year>) <volume>20</volume>:<fpage>3241</fpage>&#x2013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.1364/oe.20.003241</pub-id>
</citation>
</ref>
<ref id="B107">
<label>107.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nakajima</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Tanaka</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Hashimoto</surname>
<given-names>T</given-names>
</name>
</person-group>. <article-title>Scalable reservoir computing on coherent linear photonic processor</article-title>. <source>Commun Phys</source> (<year>2021</year>) <volume>4</volume>:<fpage>20</fpage>. <pub-id pub-id-type="doi">10.1038/s42005-021-00519-1</pub-id>
</citation>
</ref>
<ref id="B108">
<label>108.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pelucchi</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Fagas</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Aharonovich</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Englund</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Figueroa</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Gong</surname>
<given-names>Q</given-names>
</name>
<etal/>
</person-group> <article-title>The potential and global outlook of integrated photonics for quantum technologies</article-title>. <source>Nat Rev Phys</source> (<year>2022</year>) <volume>4</volume>:<fpage>194</fpage>&#x2013;<lpage>208</lpage>. <pub-id pub-id-type="doi">10.1038/s42254-021-00398-z</pub-id>
</citation>
</ref>
<ref id="B109">
<label>109.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sibson</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Kennard</surname>
<given-names>JE</given-names>
</name>
<name>
<surname>Stanisic</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Erven</surname>
<given-names>C</given-names>
</name>
<name>
<surname>O&#x2019;Brien</surname>
<given-names>JL</given-names>
</name>
<name>
<surname>Thompson</surname>
<given-names>MG</given-names>
</name>
</person-group>. <article-title>Integrated silicon photonics for high-speed quantum key distribution</article-title>. <source>Optica</source> (<year>2017</year>) <volume>4</volume>:<fpage>172</fpage>&#x2013;<lpage>7</lpage>. <pub-id pub-id-type="doi">10.1364/optica.4.000172</pub-id>
</citation>
</ref>
<ref id="B110">
<label>110.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Buck</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Coleman</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Sargsyan</surname>
<given-names>H</given-names>
</name>
</person-group>. <source>Continuous variable quantum algorithms: an introduction</source> (<year>2021</year>). <comment>arXiv preprint arXiv:2107.02151</comment>.</citation>
</ref>
<ref id="B111">
<label>111.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bunandar</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Lentine</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Long</surname>
<given-names>CM</given-names>
</name>
<name>
<surname>Boynton</surname>
<given-names>N</given-names>
</name>
<etal/>
</person-group> <article-title>Metropolitan quantum key distribution with silicon photonics</article-title>. <source>Phys Rev X</source> (<year>2018</year>) <volume>8</volume>:<fpage>021009</fpage>. <pub-id pub-id-type="doi">10.1103/physrevx.8.021009</pub-id>
</citation>
</ref>
<ref id="B112">
<label>112.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ying</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Feng</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Dhar</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Dalir</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Gu</surname>
<given-names>J</given-names>
</name>
<etal/>
</person-group> <article-title>Electronic-photonic arithmetic logic unit for high-speed computing</article-title>. <source>Nat Commun</source> (<year>2020</year>) <volume>11</volume>:<fpage>2154</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-020-16057-3</pub-id>
</citation>
</ref>
<ref id="B113">
<label>113.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ying</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Feng</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Gu</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Soref</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>DZ</given-names>
</name>
<etal/>
</person-group> <article-title>Sequential logic and pipelining in chip-based electronic-photonic digital computing</article-title>. <source>IEEE Photon J</source> (<year>2020</year>) <volume>12</volume>:<fpage>1</fpage>&#x2013;<lpage>11</lpage>. <pub-id pub-id-type="doi">10.1109/jphot.2020.3031641</pub-id>
</citation>
</ref>
<ref id="B114">
<label>114.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gostimirovic</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Ye</surname>
<given-names>WN</given-names>
</name>
</person-group>. <article-title>Ultracompact cmos-compatible optical logic using carrier depletion in microdisk resonators</article-title>. <source>Scientific Rep</source> (<year>2017</year>) <volume>7</volume>:<fpage>12603</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-017-12680-1</pub-id>
</citation>
</ref>
<ref id="B115">
<label>115.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lei</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>He</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X</given-names>
</name>
</person-group>. <article-title>Reconfigurable photonic full-adder and full-subtractor based on three-input xor gate and logic minterms</article-title>. <source>Electron Lett</source> (<year>2012</year>) <volume>48</volume>:<fpage>399</fpage>&#x2013;<lpage>400</lpage>. <pub-id pub-id-type="doi">10.1049/el.2012.0493</pub-id>
</citation>
</ref>
<ref id="B116">
<label>116.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lu</surname>
<given-names>GW</given-names>
</name>
<name>
<surname>Qin</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Ji</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Sharif</surname>
<given-names>GM</given-names>
</name>
<name>
<surname>Yamaguchi</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Flexible and re-configurable optical three-input xor logic gate of phase-modulated signals with multicast functionality for potential application in optical physical-layer network coding</article-title>. <source>Opt express</source> (<year>2016</year>) <volume>24</volume>:<fpage>2299</fpage>&#x2013;<lpage>306</lpage>. <pub-id pub-id-type="doi">10.1364/oe.24.002299</pub-id>
</citation>
</ref>
<ref id="B117">
<label>117.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ying</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Dhar</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Feng</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Mital</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Chung</surname>
<given-names>CJ</given-names>
</name>
<etal/>
</person-group> <article-title>Electro-optic ripple-carry adder in integrated silicon photonics for optical computing</article-title>. <source>IEEE J Selected Top Quan Electron</source> (<year>2018</year>) <volume>24</volume>:<fpage>1</fpage>&#x2013;<lpage>10</lpage>. <pub-id pub-id-type="doi">10.1109/JSTQE.2018.2836955</pub-id>
</citation>
</ref>
<ref id="B118">
<label>118.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rostami</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Nejad</surname>
<given-names>HBA</given-names>
</name>
<name>
<surname>Qartavol</surname>
<given-names>RM</given-names>
</name>
<name>
<surname>Saghai</surname>
<given-names>HR</given-names>
</name>
</person-group>. <article-title>Tb/s optical logic gates based on quantum-dot semiconductor optical amplifiers</article-title>. <source>IEEE J Quan Electron</source> (<year>2010</year>) <volume>46</volume>:<fpage>354</fpage>&#x2013;<lpage>60</lpage>. <pub-id pub-id-type="doi">10.1109/JQE.2009.2033253</pub-id>
</citation>
</ref>
<ref id="B119">
<label>119.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mukherjee</surname>
<given-names>K</given-names>
</name>
</person-group>. <article-title>Ultra-fast and gate using single semi-reflective quantum dot semiconductor optical amplifier</article-title>. <source>Photonic Netw Commun</source> (<year>2023</year>) <volume>45</volume>:<fpage>97</fpage>&#x2013;<lpage>106</lpage>. <pub-id pub-id-type="doi">10.1007/s11107-023-00996-0</pub-id>
</citation>
</ref>
<ref id="B120">
<label>120.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rostami</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Nejad</surname>
<given-names>HBA</given-names>
</name>
<name>
<surname>Qartavol</surname>
<given-names>RM</given-names>
</name>
<name>
<surname>Saghai</surname>
<given-names>HR</given-names>
</name>
</person-group>. <article-title>Tb/s optical logic gates based on quantum-dot semiconductor optical amplifiers</article-title>. <source>IEEE J Quan Electron</source> (<year>2010</year>) <volume>46</volume>:<fpage>354</fpage>&#x2013;<lpage>60</lpage>. <pub-id pub-id-type="doi">10.1109/jqe.2009.2033253</pub-id>
</citation>
</ref>
<ref id="B121">
<label>121.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Ye</surname>
<given-names>P</given-names>
</name>
</person-group>. <article-title>All optical xor logic gates: technologies and experiment demonstrations</article-title>. <source>IEEE Commun Mag</source> (<year>2005</year>) <volume>43</volume>:<fpage>S19</fpage>&#x2013;<lpage>S24</lpage>. <pub-id pub-id-type="doi">10.1109/mcom.2005.1453421</pub-id>
</citation>
</ref>
<ref id="B122">
<label>122.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cybenko</surname>
<given-names>G</given-names>
</name>
</person-group>. <article-title>Approximation by superpositions of a sigmoidal function</article-title>. <source>Maths Control Signals Syst</source> (<year>1989</year>) <volume>2</volume>:<fpage>303</fpage>&#x2013;<lpage>14</lpage>. <pub-id pub-id-type="doi">10.1007/BF02551274</pub-id>
</citation>
</ref>
<ref id="B123">
<label>123.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Benth</surname>
<given-names>FE</given-names>
</name>
<name>
<surname>Detering</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Galimberti</surname>
<given-names>L</given-names>
</name>
</person-group>. <article-title>Neural networks in Fr&#xe9;chet spaces</article-title>. <source>Ann Maths Artif Intelligence</source> (<year>2023</year>) <volume>91</volume>:<fpage>75</fpage>&#x2013;<lpage>103</lpage>. <pub-id pub-id-type="doi">10.1007/s10472-022-09824-z</pub-id>
</citation>
</ref>
<ref id="B124">
<label>124.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Simonyan</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Zisserman</surname>
<given-names>A</given-names>
</name>
</person-group>. <source>Very deep convolutional networks for large-scale image recognition</source> (<year>2015</year>).</citation>
</ref>
<ref id="B125">
<label>125.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Anderson</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Vasudevan</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Keane</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Gregg</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>High-performance low-memory lowering: gemm-based algorithms for dnn convolution</article-title>. In: <source>2020 IEEE 32nd international symposium on computer architecture and high performance computing (SBAC-PAD)</source> (<year>2020</year>). p. <fpage>99</fpage>&#x2013;<lpage>106</lpage>. <pub-id pub-id-type="doi">10.1109/SBAC-PAD49847.2020.00024</pub-id>
</citation>
</ref>
<ref id="B126">
<label>126.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Vasudevan</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Anderson</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Gregg</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>Parallel multi channel convolution using general matrix multiplication</article-title>. In: <source>2017 IEEE 28th international conference on application-specific systems, architectures and processors (ASAP)</source>. <publisher-loc>Los Alamitos, CA, USA</publisher-loc>: <publisher-name>IEEE Computer Society</publisher-name> (<year>2017</year>). p. <fpage>19</fpage>&#x2013;<lpage>24</lpage>. <pub-id pub-id-type="doi">10.1109/ASAP.2017.7995254</pub-id>
</citation>
</ref>
<ref id="B127">
<label>127.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shen</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Harris</surname>
<given-names>NC</given-names>
</name>
<name>
<surname>Skirlo</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Prabhu</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Baehr-Jones</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Hochberg</surname>
<given-names>M</given-names>
</name>
<etal/>
</person-group> <article-title>Deep learning with coherent nanophotonic circuits</article-title>. <source>Nat Photon</source> (<year>2017</year>) <volume>11</volume>:<fpage>441</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1109/phosst.2017.8012714</pub-id>
</citation>
</ref>
<ref id="B128">
<label>128.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shokraneh</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Geoffroy-gagnon</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Liboiron-Ladouceur</surname>
<given-names>O</given-names>
</name>
</person-group>. <article-title>The diamond mesh, a phase-error- and loss-tolerant field-programmable MZI-based optical processor for optical neural networks</article-title>. <source>Opt Express</source> (<year>2020</year>) <volume>28</volume>:<fpage>23495</fpage>&#x2013;<lpage>508</lpage>. <pub-id pub-id-type="doi">10.1364/OE.395441</pub-id>
</citation>
</ref>
<ref id="B129">
<label>129.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tait</surname>
<given-names>AN</given-names>
</name>
<name>
<surname>Nahmias</surname>
<given-names>MA</given-names>
</name>
<name>
<surname>Shastri</surname>
<given-names>BJ</given-names>
</name>
<name>
<surname>Prucnal</surname>
<given-names>PR</given-names>
</name>
</person-group>. <article-title>Broadcast and weight: an integrated network for scalable photonic spike processing</article-title>. <source>J Lightwave Tech</source> (<year>2014</year>) <volume>32</volume>:<fpage>4029</fpage>&#x2013;<lpage>41</lpage>. <pub-id pub-id-type="doi">10.1109/jlt.2014.2345652</pub-id>
</citation>
</ref>
<ref id="B130">
<label>130.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Feldmann</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Youngblood</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Karpov</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Gehring</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Stappers</surname>
<given-names>M</given-names>
</name>
<etal/>
</person-group> <article-title>Parallel convolutional processing using an integrated photonic tensor core</article-title>. <source>Nature</source> (<year>2021</year>) <volume>589</volume>:<fpage>52</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-020-03070-1</pub-id>
</citation>
</ref>
<ref id="B131">
<label>131.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xea</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Tan</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Corcoran</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Boes</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>TG</given-names>
</name>
<etal/>
</person-group> <article-title>11tops photonic convolutional accelerator for optical neural networks</article-title>. <source>Nature</source> (<year>2021</year>) <volume>589</volume>:<fpage>44</fpage>&#x2013;<lpage>51</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-020-03063-0</pub-id>
</citation>
</ref>
<ref id="B132">
<label>132.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ashtiani</surname>
<given-names>F</given-names>
</name>
<name>
<surname>On</surname>
<given-names>MB</given-names>
</name>
<name>
<surname>Sanchez-Jacome</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Perez-Lopez</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Ben Yoo</surname>
<given-names>SJ</given-names>
</name>
<name>
<surname>Blanco-Redondo</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Photonic max-pooling for deep neural networks using a programmable photonic platform</article-title>. In: <source>2023 optical fiber communications conference and exhibition (OFC)</source> (<year>2023</year>). p. <fpage>1</fpage>&#x2013;<lpage>3</lpage>. <pub-id pub-id-type="doi">10.1364/OFC.2023.M1J.6</pub-id>
</citation>
</ref>
<ref id="B133">
<label>133.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Marinis</surname>
<given-names>LD</given-names>
</name>
<name>
<surname>Nesti</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Cococcioni</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Andriolli</surname>
<given-names>N</given-names>
</name>
</person-group>. <source>A photonic accelerator for feature map generation in convolutional neural networks. OSA Advanced Photonics Congress (AP) 2020 (IPR, NP, NOMA, Networks, PVLED, PSC, SPPCom, SOF)</source>. <publisher-name>Optica Publishing Group</publisher-name> (<year>2020</year>). <pub-id pub-id-type="doi">10.1364/PSC.2020.PsTh1F.3</pub-id>
</citation>
</ref>
<ref id="B134">
<label>134.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chua</surname>
<given-names>L</given-names>
</name>
</person-group>. <article-title>Memristor-the missing circuit element</article-title>. <source>IEEE Trans Circuit Theor</source> (<year>1971</year>) <volume>18</volume>:<fpage>507</fpage>&#x2013;<lpage>19</lpage>. <pub-id pub-id-type="doi">10.1109/TCT.1971.1083337</pub-id>
</citation>
</ref>
<ref id="B135">
<label>135.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xia</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>JJ</given-names>
</name>
</person-group>. <article-title>Memristive crossbar arrays for brain-inspired computing</article-title>. <source>Nat Mater</source> (<year>2019</year>) <volume>18</volume>:<fpage>309</fpage>&#x2013;<lpage>23</lpage>. <pub-id pub-id-type="doi">10.1038/s41563-019-0291-x</pub-id>
</citation>
</ref>
<ref id="B136">
<label>136.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Shafiee</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Nag</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Muralimanohar</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Balasubramonian</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Strachan</surname>
<given-names>JP</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>M</given-names>
</name>
<etal/>
</person-group> <article-title>Isaac: a convolutional neural network accelerator with in-situ analog arithmetic in crossbars</article-title>. In: <source>2016 ACM/IEEE 43rd annual international symposium on computer architecture (ISCA)</source> (<year>2016</year>). p. <fpage>14</fpage>&#x2013;<lpage>26</lpage>.</citation>
</ref>
<ref id="B137">
<label>137.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mao</surname>
<given-names>JY</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>ST</given-names>
</name>
</person-group>. <article-title>Photonic memristor for future computing: a perspective</article-title>. <source>Adv Opt Mater</source> (<year>2019</year>) <volume>7</volume>:<fpage>1900766</fpage>. <pub-id pub-id-type="doi">10.1002/adom.201900766</pub-id>
</citation>
</ref>
<ref id="B138">
<label>138.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Choi</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Kwak</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Topologically protected all-optical memory</article-title>. <source>Adv Electron Mater</source> (<year>2022</year>) <volume>8</volume>:<fpage>2200579</fpage>. <pub-id pub-id-type="doi">10.1002/aelm.202200579</pub-id>
</citation>
</ref>
<ref id="B139">
<label>139.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nahmias</surname>
<given-names>MA</given-names>
</name>
<name>
<surname>de Lima</surname>
<given-names>TF</given-names>
</name>
<name>
<surname>Tait</surname>
<given-names>AN</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>HT</given-names>
</name>
<name>
<surname>Shastri</surname>
<given-names>BJ</given-names>
</name>
<name>
<surname>Prucnal</surname>
<given-names>PR</given-names>
</name>
</person-group>. <article-title>Photonic multiply-accumulate operations for neural networks</article-title>. <source>IEEE J Selected Top Quan Electron</source> (<year>2020</year>) <volume>26</volume>:<fpage>1</fpage>&#x2013;<lpage>18</lpage>. <pub-id pub-id-type="doi">10.1109/JSTQE.2019.2941485</pub-id>
</citation>
</ref>
<ref id="B140">
<label>140.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Miscuglio</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Sorger</surname>
<given-names>VJ</given-names>
</name>
</person-group>. <article-title>Photonic tensor cores for machine learning</article-title>. <source>Appl Phys Rev</source> (<year>2020</year>) <volume>7</volume>. <pub-id pub-id-type="doi">10.1063/5.0001942</pub-id>
</citation>
</ref>
<ref id="B141">
<label>141.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Strassen</surname>
<given-names>V</given-names>
</name>
</person-group>. <article-title>Gaussian elimination is not optimal</article-title>. <source>Numerische Mathematik</source> (<year>1969</year>) <volume>13</volume>:<fpage>354</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1007/BF02165411</pub-id>
</citation>
</ref>
<ref id="B142">
<label>142.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Coppersmith</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Winograd</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Matrix multiplication via arithmetic progressions</article-title>. <source>J Symbolic Comput</source> (<year>1990</year>) <volume>9</volume>:<fpage>251</fpage>&#x2013;<lpage>80</lpage>. <pub-id pub-id-type="doi">10.1016/S0747-7171(08)80013-2</pub-id>
</citation>
</ref>
<ref id="B143">
<label>143.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Shiflett</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Karanth</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Bunescu</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Louri</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Albireo: energy-efficient acceleration of convolutional neural networks via silicon photonics</article-title>. In: <source>2021 ACM/IEEE 48th annual international symposium on computer architecture (ISCA) (IEEE)</source> (<year>2021</year>). p. <fpage>860</fpage>&#x2013;<lpage>73</lpage>.</citation>
</ref>
<ref id="B144">
<label>144.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Shiflett</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Wright</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Karanth</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Louri</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Pixel: photonic neural network accelerator</article-title>. In: <source>2020 IEEE international symposium on high performance computer architecture (HPCA) (IEEE)</source> (<year>2020</year>). p. <fpage>474</fpage>&#x2013;<lpage>87</lpage>.</citation>
</ref>
<ref id="B145">
<label>145.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Peng</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Alkabani</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Puri</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Sorger</surname>
<given-names>V</given-names>
</name>
<name>
<surname>El-Ghazawi</surname>
<given-names>T</given-names>
</name>
</person-group>. <article-title>A deep neural network accelerator using residue arithmetic in a hybrid optoelectronic system</article-title>. <source>ACM J Emerging Tech Comput Syst (Jetc)</source> (<year>2022</year>) <volume>18</volume>:<fpage>1</fpage>&#x2013;<lpage>26</lpage>. <pub-id pub-id-type="doi">10.1145/3550273</pub-id>
</citation>
</ref>
<ref id="B146">
<label>146.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Dang</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Dass</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Mahapatra</surname>
<given-names>R</given-names>
</name>
</person-group>. <article-title>Convlight: a convolutional accelerator with memristor integrated photonic computing</article-title>. In: <source>2017 IEEE 24th international conference on high performance computing (HiPC)</source>. <publisher-name>IEEE</publisher-name> (<year>2017</year>). p. <fpage>114</fpage>&#x2013;<lpage>23</lpage>.</citation>
</ref>
<ref id="B147">
<label>147.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Deng</surname>
<given-names>L</given-names>
</name>
</person-group>. <article-title>The mnist database of handwritten digit images for machine learning research</article-title>. <source>IEEE Signal Process. Mag</source> (<year>2012</year>) <volume>29</volume>:<fpage>141</fpage>&#x2013;<lpage>2</lpage>.</citation>
</ref>
<ref id="B148">
<label>148.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Mehrabian</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Al-Kabani</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Sorger</surname>
<given-names>VJ</given-names>
</name>
<name>
<surname>El-Ghazawi</surname>
<given-names>T</given-names>
</name>
</person-group>. <article-title>Pcnna: a photonic convolutional neural network accelerator</article-title>. In: <source>2018 31st IEEE international system-on-chip conference (SOCC) (IEEE)</source> (<year>2018</year>). p. <fpage>169</fpage>&#x2013;<lpage>73</lpage>.</citation>
</ref>
<ref id="B149">
<label>149.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sunny</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Mirza</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Nikdast</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Pasricha</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Crosslight: a cross-layer optimized silicon photonic neural network accelerator</article-title>. In: <source>2021 58th ACM/IEEE design automation conference (DAC)</source> (<year>2021</year>). p. <fpage>1069</fpage>&#x2013;<lpage>74</lpage>. <pub-id pub-id-type="doi">10.1109/DAC18074.2021.9586161</pub-id>
</citation>
</ref>
<ref id="B150">
<label>150.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sunny</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Nikdast</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Pasricha</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Sonic: a sparse neural network inference accelerator with silicon photonics for energy-efficient deep learning</article-title>. In: <source>2022 27th asia and south pacific design automation conference (ASP-DAC)</source> (<year>2022</year>). p. <fpage>214</fpage>&#x2013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.1109/ASP-DAC52403.2022.9712530</pub-id>
</citation>
</ref>
<ref id="B151">
<label>151.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Han</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Mao</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Dally</surname>
<given-names>WJ</given-names>
</name>
</person-group>. <article-title>Deep compression: compressing deep neural networks with pruning</article-title>. In: <source>Trained quantization and huffman coding</source> (<year>2016</year>). <pub-id pub-id-type="doi">10.48550/arXiv.1510.00149</pub-id>
</citation>
</ref>
<ref id="B152">
<label>152.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zhu</surname>
<given-names>MH</given-names>
</name>
<name>
<surname>Gupta</surname>
<given-names>S</given-names>
</name>
</person-group>. <source>To prune, or not to prune: exploring the efficacy of pruning for model compression</source> (<year>2018</year>).</citation>
</ref>
<ref id="B153">
<label>153.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Y</given-names>
</name>
</person-group>. <article-title>Incremental network quantization: towards lossless CNNs with low-precision weights</article-title>. In: <source>International conference on learning representations</source> (<year>2017</year>).</citation>
</ref>
<ref id="B154">
<label>154.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Judd</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Albericio</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Hetherington</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Aamodt</surname>
<given-names>TM</given-names>
</name>
<name>
<surname>Moshovos</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Stripes: bit-serial deep neural network computing</article-title>. In: <source>2016 49th annual IEEE/ACM international symposium on microarchitecture (MICRO)</source> (<year>2016</year>). p. <fpage>1</fpage>&#x2013;<lpage>12</lpage>. <pub-id pub-id-type="doi">10.1109/MICRO.2016.7783722</pub-id>
</citation>
</ref>
<ref id="B155">
<label>155.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Shiflett</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Karanth</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Louri</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Bunescu</surname>
<given-names>R</given-names>
</name>
</person-group>. <article-title>Bitwise neural network acceleration using silicon photonics</article-title>. In: <source>Proceedings of the 2021 on great lakes symposium on VLSI</source> (<year>2021</year>). p. <fpage>9</fpage>&#x2013;<lpage>14</lpage>.</citation>
</ref>
<ref id="B156">
<label>156.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zokaee</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Lou</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Youngblood</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>L</given-names>
</name>
</person-group>. <article-title>Lightbulb: a photonic-nonvolatile-memory-based accelerator for binarized convolutional neural networks</article-title>. In: <source>2020 design, automation &#x26; test in europe conference &#x26; exhibition (DATE) (IEEE)</source> (<year>2020</year>). p. <fpage>1438</fpage>&#x2013;<lpage>43</lpage>.</citation>
</ref>
<ref id="B157">
<label>157.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Danial</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Wainstein</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Kraus</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Kvatinsky</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Breaking through the speed-power-accuracy tradeoff in ADCs using a memristive neuromorphic architecture</article-title>. <source>IEEE Trans Emerging Top Comput Intelligence</source> (<year>2018</year>) <volume>2</volume>:<fpage>396</fpage>&#x2013;<lpage>409</lpage>. <pub-id pub-id-type="doi">10.1109/TETCI.2018.2849109</pub-id>
</citation>
</ref>
<ref id="B158">
<label>158.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sunny</surname>
<given-names>FP</given-names>
</name>
<name>
<surname>Mirza</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Nikdast</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Pasricha</surname>
<given-names>S</given-names>
</name>
</person-group>. <source>Robin: a robust optical binary neural network accelerator</source> (<year>2021</year>). <pub-id pub-id-type="doi">10.1145/3476988</pub-id>
</citation>
</ref>
<ref id="B159">
<label>159.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sunny</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Nikdast</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Pasricha</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>A silicon photonic accelerator for convolutional neural networks with heterogeneous quantization</article-title>. In: <source>Proceedings of the great lakes symposium on VLSI 2022</source> (<year>2022</year>). p. <fpage>367</fpage>&#x2013;<lpage>71</lpage>.</citation>
</ref>
<ref id="B160">
<label>160.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Peng</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Alkabani</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Sorger</surname>
<given-names>VJ</given-names>
</name>
<name>
<surname>El-Ghazawi</surname>
<given-names>T</given-names>
</name>
</person-group>. <article-title>Dnnara: a deep neural network accelerator using residue arithmetic and integrated photonics</article-title>. In: <source>Proceedings of the 49th international conference on parallel processing</source> (<year>2020</year>). p. <fpage>1</fpage>&#x2013;<lpage>11</lpage>.</citation>
</ref>
<ref id="B161">
<label>161.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Dosovitskiy</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Beyer</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Kolesnikov</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Weissenborn</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Zhai</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Unterthiner</surname>
<given-names>T</given-names>
</name>
<etal/>
</person-group> <source>An image is worth 16x16 words: Transformers for image recognition at scale</source> (<year>2020</year>). <comment>CoRR abs/2010.11929</comment>.</citation>
</ref>
<ref id="B162">
<label>162.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Devlin</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Chang</surname>
<given-names>MW</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Toutanova</surname>
<given-names>K</given-names>
</name>
</person-group>. <article-title>BERT: pre-training of deep bidirectional Transformers for language understanding</article-title>. In: <person-group person-group-type="editor">
<name>
<surname>Burstein</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Doran</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Solorio</surname>
<given-names>T</given-names>
</name>
</person-group>, editors. <source>Proceedings of the 2019 conference of the north American chapter of the association for computational linguistics: human language technologies, volume 1 (long and Short papers)</source>. <publisher-loc>Minneapolis, Minnesota</publisher-loc>: <publisher-name>Association for Computational Linguistics</publisher-name> (<year>2019</year>). p. <fpage>4171</fpage>&#x2013;<lpage>86</lpage>. <pub-id pub-id-type="doi">10.18653/v1/N19-1423</pub-id>
</citation>
</ref>
<ref id="B163">
<label>163.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Afifi</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Sunny</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Nikdast</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Pasricha</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Tron: Transformer neural network acceleration with non-coherent silicon photonics</article-title>. In: <source>Proceedings of the great lakes symposium on VLSI 2023</source> (<year>2023</year>). p. <fpage>15</fpage>&#x2013;<lpage>21</lpage>.</citation>
</ref>
<ref id="B164">
<label>164.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sunny</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Nikdast</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Pasricha</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Reclight: a recurrent neural network accelerator with integrated silicon photonics</article-title>. In: <source>2022 IEEE computer society annual symposium on VLSI (ISVLSI) (IEEE)</source> (<year>2022</year>). p. <fpage>98</fpage>&#x2013;<lpage>103</lpage>.</citation>
</ref>
<ref id="B165">
<label>165.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hochreiter</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Schmidhuber</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Lstm can solve hard long time lag problems</article-title>. <source>Adv Neural Inf Process Syst</source> (<year>1996</year>) <volume>9</volume>.</citation>
</ref>
<ref id="B166">
<label>166.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sarantoglou</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Bogris</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Mesaritakis</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Theodoridis</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Bayesian photonic accelerators for energy efficient and noise robust neural processing</article-title>. <source>IEEE J Selected Top Quan Electron</source> (<year>2022</year>) <volume>28</volume>:<fpage>1</fpage>&#x2013;<lpage>10</lpage>. <pub-id pub-id-type="doi">10.1109/JSTQE.2022.3183444</pub-id>
</citation>
</ref>
<ref id="B167">
<label>167.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>P&#xe9;rez-L&#xf3;pez</surname>
<given-names>D</given-names>
</name>
<name>
<surname>L&#xf3;pez</surname>
<given-names>A</given-names>
</name>
<name>
<surname>DasMahapatra</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Capmany</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Multipurpose self-configuration of programmable photonic circuits</article-title>. <source>Nat Commun</source> (<year>2020</year>) <volume>11</volume>:<fpage>6359</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-020-19608-w</pub-id>
</citation>
</ref>
<ref id="B168">
<label>168.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Demirkiran</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Eris</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Elmhurst</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Moore</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Harris</surname>
<given-names>NC</given-names>
</name>
<etal/>
</person-group> <article-title>An electro-photonic system for accelerating deep neural networks</article-title>. In: <source>ACM journal on emerging technologies in computing systems 19</source> (<year>2023</year>). <pub-id pub-id-type="doi">10.1145/3606949</pub-id>
</citation>
</ref>
<ref id="B169">
<label>169.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Ren</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J</given-names>
</name>
</person-group>. <source>Deep residual learning for image recognition</source> (<year>2015</year>). <comment>CoRR abs/1512.03385</comment>.</citation>
</ref>
<ref id="B170">
<label>170.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Sainath</surname>
<given-names>TN</given-names>
</name>
<name>
<surname>Prabhavalkar</surname>
<given-names>R</given-names>
</name>
<name>
<surname>McGraw</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Alvarez</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>D</given-names>
</name>
<etal/>
</person-group> <source>Streaming end-to-end speech recognition for mobile devices (arXiv)</source> (<year>2018</year>). <pub-id pub-id-type="doi">10.48550/arXiv.1811.06621</pub-id>
</citation>
</ref>
<ref id="B171">
<label>171.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Zheng</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Louri</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Karanth</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Ascend: a scalable and energy-efficient deep neural network accelerator with photonic interconnects</article-title>. <source>IEEE Trans Circuits Syst Regular Pap</source> (<year>2022</year>) <volume>69</volume>:<fpage>2730</fpage>&#x2013;<lpage>41</lpage>. <pub-id pub-id-type="doi">10.1109/TCSI.2022.3169953</pub-id>
</citation>
</ref>
<ref id="B172">
<label>172.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Narayan</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Thonnart</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Vivet</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Coskun</surname>
<given-names>AK</given-names>
</name>
</person-group>. <article-title>Prowaves: proactive runtime wavelength selection for energy-efficient photonic nocs</article-title>. <source>IEEE Trans Computer-Aided Des Integrated Circuits Syst</source> (<year>2020</year>) <volume>40</volume>:<fpage>2156</fpage>&#x2013;<lpage>69</lpage>. <pub-id pub-id-type="doi">10.1109/tcad.2020.3037327</pub-id>
</citation>
</ref>
<ref id="B173">
<label>173.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vantrease</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Schreiber</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Monchiero</surname>
<given-names>M</given-names>
</name>
<name>
<surname>McLaren</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Jouppi</surname>
<given-names>NP</given-names>
</name>
<name>
<surname>Fiorentino</surname>
<given-names>M</given-names>
</name>
<etal/>
</person-group> <article-title>Corona: system implications of emerging nanophotonic technology</article-title>. <source>ACM SIGARCH Comput Architecture News</source> (<year>2008</year>) <volume>36</volume>:<fpage>153</fpage>&#x2013;<lpage>64</lpage>. <pub-id pub-id-type="doi">10.1109/isca.2008.35</pub-id>
</citation>
</ref>
<ref id="B174">
<label>174.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sludds</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Bandyopadhyay</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Zhong</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Cochrane</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Bernstein</surname>
<given-names>L</given-names>
</name>
<etal/>
</person-group> <article-title>Delocalized photonic deep learning on the internet&#x2019;s edge</article-title>. <source>Science</source> (<year>2022</year>) <volume>378</volume>:<fpage>270</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1126/science.abq8271</pub-id>
</citation>
</ref>
<ref id="B175">
<label>175.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Giamougiannis</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Tsakyridis</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Moralis-Pegios</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Mourgias-Alexandris</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Totovic</surname>
<given-names>AR</given-names>
</name>
<name>
<surname>Dabos</surname>
<given-names>G</given-names>
</name>
<etal/>
</person-group> <article-title>Neuromorphic silicon photonics with 50 ghz tiled matrix multiplication for deep-learning applications</article-title>. <source>Adv Photon</source> (<year>2023</year>) <volume>5</volume>:<fpage>016004</fpage>. <pub-id pub-id-type="doi">10.1117/1.ap.5.1.016004</pub-id>
</citation>
</ref>
<ref id="B176">
<label>176.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Lou</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>L</given-names>
</name>
</person-group>. <article-title>Mindreading: an ultra-low-power photonic accelerator for eeg-based human intention recognition</article-title>. In: <source>2020 25th asia and south pacific design automation conference</source>. <publisher-name>ASP-DAC</publisher-name> (<year>2020</year>). p. <fpage>464</fpage>&#x2013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.1109/ASP-DAC47756.2020.9045333</pub-id>
</citation>
</ref>
<ref id="B177">
<label>177.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Midolo</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Schliesser</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Fiore</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Nano-opto-electro-mechanical systems</article-title>. <source>Nat nanotechnology</source> (<year>2018</year>) <volume>13</volume>:<fpage>11</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1038/s41565-017-0039-1</pub-id>
</citation>
</ref>
<ref id="B178">
<label>178.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ki</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Notomi</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Naruse</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Inoue</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Kawakami</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Uchida</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Novel frontier of photonics for data processing&#x2014;photonic accelerator</article-title>. <source>Apl Photon</source> (<year>2019</year>) <volume>4</volume>. <pub-id pub-id-type="doi">10.1063/1.5108912</pub-id>
</citation>
</ref>
<ref id="B179">
<label>179.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shafiee</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Banerjee</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Chakrabarty</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Pasricha</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Nikdast</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>Analysis of optical loss and crosstalk noise in MZI-based coherent photonic neural networks</article-title>. <source>J Lightwave Tech</source> (<year>2024</year>) <fpage>1</fpage>&#x2013;<lpage>16</lpage>. <pub-id pub-id-type="doi">10.1109/JLT.2024.3373250</pub-id>
</citation>
</ref>
<ref id="B180">
<label>180.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yu</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>N</given-names>
</name>
</person-group>. <article-title>Heavy tails and pruning in programmable photonic circuits for universal unitaries</article-title>. <source>Nat Commun</source> (<year>2023</year>) <volume>14</volume>:<fpage>1853</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-023-37611-9</pub-id>
</citation>
</ref>
<ref id="B181">
<label>181.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Buddhiraju</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Dutt</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Minkov</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Williamson</surname>
<given-names>IAD</given-names>
</name>
<name>
<surname>Fan</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Arbitrary linear transformations for photons in the frequency synthetic dimension</article-title>. <source>Nat Commun</source> (<year>2021</year>) <volume>12</volume>:<fpage>2401</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-021-22670-7</pub-id>
</citation>
</ref>
<ref id="B182">
<label>182.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Piao</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>N</given-names>
</name>
</person-group>. <source>Programmable photonic time circuits for highly scalable universal unitaries</source> (<year>2023</year>). <pub-id pub-id-type="doi">10.48550/arXiv.2305.17632</pub-id>
</citation>
</ref>
<ref id="B183">
<label>183.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bandyopadhyay</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Hamerly</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Englund</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>Hardware error correction for programmable photonics</article-title>. <source>Optica</source> (<year>2021</year>) <volume>8</volume>:<fpage>1247</fpage>&#x2013;<lpage>55</lpage>. <pub-id pub-id-type="doi">10.1364/OPTICA.424052</pub-id>
</citation>
</ref>
<ref id="B184">
<label>184.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>ZD</given-names>
</name>
<name>
<surname>Tang</surname>
<given-names>JD</given-names>
</name>
<name>
<surname>Lv</surname>
<given-names>XM</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>ML</given-names>
</name>
<etal/>
</person-group> <article-title>Recent advances in nano-opto-electro-mechanical systems</article-title>. <source>Nanophotonics</source> (<year>2021</year>) <volume>10</volume>:<fpage>2265</fpage>&#x2013;<lpage>81</lpage>. <pub-id pub-id-type="doi">10.1515/nanoph-2021-0082</pub-id>
</citation>
</ref>
<ref id="B185">
<label>185.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shakoor</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Nozaki</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Kuramochi</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Nishiguchi</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Shinya</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Notomi</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>Compact 1d-silicon photonic crystal electro-optic modulator operating with ultra-low switching voltage and energy</article-title>. <source>Opt express</source> (<year>2014</year>) <volume>22</volume>:<fpage>28623</fpage>&#x2013;<lpage>34</lpage>. <pub-id pub-id-type="doi">10.1364/oe.22.028623</pub-id>
</citation>
</ref>
<ref id="B186">
<label>186.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kim</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>JW</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>IG</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>JM</given-names>
</name>
<etal/>
</person-group> <article-title>Low-voltage high-performance silicon photonic devices and photonic integrated circuits operating up to 30 gb/s</article-title>. <source>Opt Express</source> (<year>2011</year>) <volume>19</volume>:<fpage>26936</fpage>&#x2013;<lpage>47</lpage>. <pub-id pub-id-type="doi">10.1364/oe.19.026936</pub-id>
</citation>
</ref>
<ref id="B187">
<label>187.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jayatilleka</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Shoman</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Chrostowski</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Shekhar</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Photoconductive heaters enable control of large-scale silicon photonic ring resonator circuits</article-title>. <source>Optica</source> (<year>2019</year>) <volume>6</volume>:<fpage>84</fpage>&#x2013;<lpage>91</lpage>. <pub-id pub-id-type="doi">10.1364/optica.6.000084</pub-id>
</citation>
</ref>
<ref id="B188">
<label>188.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Buckley</surname>
<given-names>SM</given-names>
</name>
<name>
<surname>Tait</surname>
<given-names>AN</given-names>
</name>
<name>
<surname>McCaughan</surname>
<given-names>AN</given-names>
</name>
<name>
<surname>Shastri</surname>
<given-names>BJ</given-names>
</name>
</person-group>. <article-title>Photonic online learning: a perspective</article-title>. <source>Nanophotonics</source> (<year>2023</year>) <volume>12</volume>:<fpage>833</fpage>&#x2013;<lpage>45</lpage>. <pub-id pub-id-type="doi">10.1515/nanoph-2022-0553</pub-id>
</citation>
</ref>
<ref id="B189">
<label>189.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pai</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Hughes</surname>
<given-names>TW</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Bartlett</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Williamson</surname>
<given-names>IAD</given-names>
</name>
<etal/>
</person-group> <article-title>Experimentally realized <italic>in situ</italic> backpropagation for deep learning in photonic neural networks</article-title>. <source>Science</source> (<year>2023</year>) <volume>380</volume>:<fpage>398</fpage>&#x2013;<lpage>404</lpage>. <pub-id pub-id-type="doi">10.1126/science.ade8450</pub-id>
</citation>
</ref>
<ref id="B190">
<label>190.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Spall</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Lvovsky</surname>
<given-names>AI</given-names>
</name>
</person-group>. <article-title>Hybrid training of optical neural networks</article-title>. <source>Optica</source> (<year>2022</year>) <volume>9</volume>:<fpage>803</fpage>&#x2013;<lpage>11</lpage>. <pub-id pub-id-type="doi">10.1364/fio.2022.ftu6d.2</pub-id>
</citation>
</ref>
<ref id="B191">
<label>191.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Spall</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Lvovsky</surname>
<given-names>AI</given-names>
</name>
</person-group>. <source>Training neural networks with end-to-end optical backpropagation</source> (<year>2023</year>). <pub-id pub-id-type="doi">10.48550/arXiv.2308.05226</pub-id>
</citation>
</ref>
<ref id="B192">
<label>192.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dang</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Chittamuru</surname>
<given-names>SVR</given-names>
</name>
<name>
<surname>Pasricha</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Mahapatra</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Sahoo</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>Bplight-cnn: a photonics-based backpropagation accelerator for deep learning</article-title>. <source>ACM J Emerging Tech Comput Syst (Jetc)</source> (<year>2021</year>) <volume>17</volume>:<fpage>1</fpage>&#x2013;<lpage>26</lpage>. <pub-id pub-id-type="doi">10.1145/3446212</pub-id>
</citation>
</ref>
<ref id="B193">
<label>193.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dang</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Sahoo</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>Litecon: an all-photonic neuromorphic accelerator for energy-efficient deep learning</article-title>. <source>ACM Trans Architecture Code Optimization (Taco)</source> (<year>2022</year>) <volume>19</volume>:<fpage>1</fpage>&#x2013;<lpage>22</lpage>. <pub-id pub-id-type="doi">10.1145/3531226</pub-id>
</citation>
</ref>
<ref id="B194">
<label>194.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Bandyopadhyay</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Sludds</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Krastanov</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Hamerly</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Harris</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Bunandar</surname>
<given-names>D</given-names>
</name>
<etal/>
</person-group> <article-title>A photonic deep neural network processor on a single chip with optically accelerated training</article-title>. In: <source>Cleo 2023</source>. <publisher-name>Optica Publishing Group</publisher-name> (<year>2023</year>). <comment>SM2P.2</comment>. <pub-id pub-id-type="doi">10.1364/CLEO_SI.2023.SM2P.2</pub-id>
</citation>
</ref>
<ref id="B195">
<label>195.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kovachki</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Azizzadenesheli</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Bhattacharya</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Stuart</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Neural operator: learning maps between function spaces with applications to PDEs</article-title>. <source>J Machine Learn Res</source> (<year>2023</year>) <volume>24</volume>.</citation>
</ref>
<ref id="B196">
<label>196.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Ohana</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Wacker</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Marmin</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Krzakala</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Filippone</surname>
<given-names>M</given-names>
</name>
<etal/>
</person-group> <article-title>Kernel computations from large-scale random features obtained by optical processing units</article-title>. In: <source>ICASSP 2020-2020 IEEE international conference on acoustics, speech and signal processing (ICASSP)</source>. <publisher-name>IEEE</publisher-name> (<year>2020</year>). p. <fpage>9294</fpage>&#x2013;<lpage>8</lpage>.</citation>
</ref>
<ref id="B197">
<label>197.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Gu</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>XD</given-names>
</name>
<name>
<surname>Thompson</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Paesani</surname>
<given-names>S</given-names>
</name>
<etal/>
</person-group> <article-title>An optical neural chip for implementing complex-valued neural network</article-title>. <source>Nat Commun</source> (<year>2021</year>) <volume>12</volume>:<fpage>457</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-020-20719-7</pub-id>
</citation>
</ref>
<ref id="B198">
<label>198.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nikkhah</surname>
<given-names>V</given-names>
</name>
<name>
<surname>Mencagli</surname>
<given-names>MJ</given-names>
</name>
<name>
<surname>Engheta</surname>
<given-names>N</given-names>
</name>
</person-group>. <article-title>Reconfigurable nonlinear optical element using tunable couplers and inverse-designed structure</article-title>. <source>Nanophotonics</source> (<year>2023</year>) <volume>12</volume>:<fpage>3019</fpage>&#x2013;<lpage>27</lpage>. <pub-id pub-id-type="doi">10.1515/nanoph-2023-0152</pub-id>
</citation>
</ref>
<ref id="B199">
<label>199.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Liao</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Su</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Geng</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J</given-names>
</name>
<etal/>
</person-group> <article-title>Tunable on-chip mode converter enabled by inverse design</article-title>. <source>Nanophotonics</source> (<year>2023</year>) <volume>12</volume>:<fpage>1105</fpage>&#x2013;<lpage>14</lpage>. <pub-id pub-id-type="doi">10.1515/nanoph-2022-0638</pub-id>
</citation>
</ref>
<ref id="B200">
<label>200.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pan</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>X</given-names>
</name>
</person-group>. <article-title>Deep learning and adjoint method accelerated inverse design in photonics: a review</article-title>. <source>Photonics</source> (<year>2023</year>) <volume>10</volume>:<fpage>852</fpage>. <pub-id pub-id-type="doi">10.3390/photonics10070852</pub-id>
</citation>
</ref>
<ref id="B201">
<label>201.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Park</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Nam</surname>
<given-names>DW</given-names>
</name>
<name>
<surname>Chung</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>CY</given-names>
</name>
<name>
<surname>Jang</surname>
<given-names>MS</given-names>
</name>
</person-group>. <article-title>Free-form optimization of nanophotonic devices: from classical methods to deep learning</article-title>. <source>Nanophotonics</source> (<year>2022</year>) <volume>11</volume>:<fpage>1809</fpage>&#x2013;<lpage>45</lpage>. <pub-id pub-id-type="doi">10.1515/nanoph-2021-0713</pub-id>
</citation>
</ref>
<ref id="B202">
<label>202.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sanz</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Lamata</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Solano</surname>
<given-names>E</given-names>
</name>
</person-group>. <article-title>Invited article: quantum memristors in quantum photonics</article-title>. <source>APL Photon</source> (<year>2018</year>) <volume>3</volume>:<fpage>080801</fpage>. <pub-id pub-id-type="doi">10.1063/1.5036596</pub-id>
</citation>
</ref>
<ref id="B203">
<label>203.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Spagnolo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Morris</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Piacentini</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Antesberger</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Massa</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Crespi</surname>
<given-names>A</given-names>
</name>
<etal/>
</person-group> <article-title>Experimental photonic quantum memristor</article-title>. <source>Nat Photon</source> (<year>2022</year>) <volume>16</volume>:<fpage>318</fpage>&#x2013;<lpage>23</lpage>. <pub-id pub-id-type="doi">10.1038/s41566-022-00973-5</pub-id>
</citation>
</ref>
<ref id="B204">
<label>204.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Steinbrecher</surname>
<given-names>GR</given-names>
</name>
<name>
<surname>Olson</surname>
<given-names>JP</given-names>
</name>
<name>
<surname>Englund</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Carolan</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Quantum optical neural networks</article-title>. <source>npj Quan Inf</source> (<year>2019</year>) <volume>5</volume>:<fpage>60</fpage>. <pub-id pub-id-type="doi">10.1038/s41534-019-0174-7</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>