<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Phys.</journal-id>
<journal-title>Frontiers in Physics</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Phys.</abbrev-journal-title>
<issn pub-type="epub">2296-424X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1529188</article-id>
<article-id pub-id-type="doi">10.3389/fphy.2025.1529188</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Physics</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Optimizing quantum convolutional neural network architectures for arbitrary data dimension</article-title>
<alt-title alt-title-type="left-running-head">Lee et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fphy.2025.1529188">10.3389/fphy.2025.1529188</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Lee</surname>
<given-names>Changwon</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2939943/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Araujo</surname>
<given-names>Israel F.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Kim</surname>
<given-names>Dongha</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Lee</surname>
<given-names>Junghan</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2977635/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Park</surname>
<given-names>Siheon</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Ryu</surname>
<given-names>Ju-Young</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2898425/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Park</surname>
<given-names>Daniel K.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2820653/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Department of Statistics and Data Science</institution>, <institution>Yonsei University</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>School of Electrical Engineering</institution>, <institution>Korea Advanced Institute of Science and Technology (KAIST)</institution>, <addr-line>Daejeon</addr-line>, <country>Republic of Korea</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Department of Physics</institution>, <institution>Yonsei University</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Department of Physics and Astronomy</institution>, <institution>Seoul National University</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Quantum AI Team</institution>, <institution>Norma Inc.</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country>
</aff>
<aff id="aff6">
<sup>6</sup>
<institution>Department of Applied Statistics</institution>, <institution>Yonsei University</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2553883/overview">Jaewoo Joo</ext-link>, University of Portsmouth, United Kingdom</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2116506/overview">De-Sheng Li</ext-link>, Hunan Institute of Engineering, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2931820/overview">Robson Christie</ext-link>, University of Portsmouth, United Kingdom</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Daniel K. Park, <email>dkd.park@yonsei.ac.kr</email>
</corresp>
<fn fn-type="equal" id="fn001">
<label>
<sup>&#x2020;</sup>
</label>
<p>These authors have contributed equally to this work</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>03</day>
<month>03</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>13</volume>
<elocation-id>1529188</elocation-id>
<history>
<date date-type="received">
<day>16</day>
<month>11</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>03</day>
<month>02</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2025 Lee, Araujo, Kim, Lee, Park, Ryu and Park.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Lee, Araujo, Kim, Lee, Park, Ryu and Park</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Quantum convolutional neural networks (QCNNs) represent a promising approach in quantum machine learning, paving new directions for both quantum and classical data analysis. This approach is particularly attractive due to the absence of the barren plateau problem, a fundamental challenge in training quantum neural networks (QNNs), and its feasibility. However, a limitation arises when applying QCNNs to classical data. The network architecture is most natural when the number of input qubits is a power of two, as this number is reduced by a factor of two in each pooling layer. The number of input qubits determines the dimensions (i.e., the number of features) of the input data that can be processed, restricting the applicability of QCNN algorithms to real-world data. To address this issue, we propose a QCNN architecture capable of handling arbitrary input data dimensions while optimizing the allocation of quantum resources such as ancillary qubits and quantum gates. This optimization is not only important for minimizing computational resources, but also essential in noisy intermediate-scale quantum (NISQ) computing, as the size of the quantum circuits that can be executed reliably is limited. Through numerical simulations, we benchmarked the classification performance of various QCNN architectures across multiple datasets with arbitrary input data dimensions, including MNIST, Landsat satellite, Fashion-MNIST, and Ionosphere. The results validate that the proposed QCNN architecture achieves excellent classification performance while utilizing a minimal resource overhead, providing an optimal solution when reliable quantum computation is constrained by noise and imperfections.</p>
</abstract>
<kwd-group>
<kwd>quantum computing</kwd>
<kwd>quantum machine learning</kwd>
<kwd>machine learning</kwd>
<kwd>quantum circuit</kwd>
<kwd>quantum algorithm</kwd>
</kwd-group>
<contract-num rid="cn001">2019-0-00003</contract-num>
<contract-num rid="cn002">2022M3E4A1074591 2023M3K5A1094813</contract-num>
<contract-num rid="cn003">2024-22-0147</contract-num>
<contract-num rid="cn004">2E32941-24-008</contract-num>
<contract-sponsor id="cn001">Institute for Information and Communications Technology Promotion<named-content content-type="fundref-id">10.13039/501100010418</named-content>
</contract-sponsor>
<contract-sponsor id="cn002">National Research Foundation of Korea<named-content content-type="fundref-id">10.13039/501100003725</named-content>
</contract-sponsor>
<contract-sponsor id="cn003">Yonsei University<named-content content-type="fundref-id">10.13039/501100002573</named-content>
</contract-sponsor>
<contract-sponsor id="cn004">Korea Institute of Science and Technology<named-content content-type="fundref-id">10.13039/501100003693</named-content>
</contract-sponsor>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Quantum Engineering and Technology</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>The advent of deep neural networks (DNNs) has transformed machine learning, drawing considerable research attention owing to the efficacy and broad applicability of DNNs [<xref ref-type="bibr" rid="B1">1</xref>, <xref ref-type="bibr" rid="B2">2</xref>]. Among the DNNs, convolutional neural networks (CNNs) have emerged to pivotally contribute toward image processing and vision tasks [<xref ref-type="bibr" rid="B3">3</xref>, <xref ref-type="bibr" rid="B4">4</xref>]. By leveraging filtering techniques, the CNN architecture effectively detects and extracts spatial features from input data. CNNs exhibit exceptional performance in diverse domains&#x2014;including image classification, object detection, face recognition, and medical image processing&#x2014;and have attracted interest from both researchers and industry [<xref ref-type="bibr" rid="B5">5</xref>&#x2013;<xref ref-type="bibr" rid="B9">9</xref>].</p>
<p>Although DNNs have proven successful in various data analytics tasks, the increasing volume and complexity of datasets present a challenge to the current classical computing paradigm, prompting the exploration of alternative solutions. Quantum machine learning (QML) has emerged as a promising approach to address the fundamental limitations of classical machine learning. By leveraging the advantages of quantum computing techniques and algorithms, QML aims to overcome the inherent constraints of its classical counterparts [<xref ref-type="bibr" rid="B10">10</xref>&#x2013;<xref ref-type="bibr" rid="B13">13</xref>]. However, a challenge in contemporary quantum computing lies in the difficulty of constructing quantum hardware. This challenge is characterized by noisy intermediate-scale quantum (NISQ) computing [<xref ref-type="bibr" rid="B14">14</xref>, <xref ref-type="bibr" rid="B15">15</xref>], as the number of quantum processors that can be controlled reliably is limited owing to noise. Quantum-classical hybrid approaches based on parameterized quantum circuits (PQCs) have been developed to enhance the utility of NISQ devices [<xref ref-type="bibr" rid="B16">16</xref>&#x2013;<xref ref-type="bibr" rid="B18">18</xref>]. These strategies have contributed to advancements in quantum computing and machine learning, facilitating improved performance and applicability in various domains. In particular, PQC-based QML models have demonstrated a potential to outperform classical models in terms of sample complexity, generalization, and trainability [<xref ref-type="bibr" rid="B19">19</xref>&#x2013;<xref ref-type="bibr" rid="B25">25</xref>]. However, PQCs encounter a critical challenge in addressing real-world problems, particularly in relation to scalability, which is attributed to a phenomenon known as barren plateaus (BP) [<xref ref-type="bibr" rid="B26">26</xref>, <xref ref-type="bibr" rid="B27">27</xref>]. This phenomenon is characterized by an intrinsic tradeoff between the expressibility and trainability of PQCs [<xref ref-type="bibr" rid="B28">28</xref>], causing the gradient of the cost function to vanish exponentially with the number of qubits under certain conditions. An effective strategy for avoiding BPs is to adopt a hierarchical quantum circuit structure, wherein the number of qubits decreases exponentially with the depth of the quantum circuit [<xref ref-type="bibr" rid="B29">29</xref>, <xref ref-type="bibr" rid="B30">30</xref>]. Quantum convolutional neural networks (QCNNs) notably employ this strategy, as highlighted in recent studies [<xref ref-type="bibr" rid="B31">31</xref>&#x2013;<xref ref-type="bibr" rid="B37">37</xref>]. Inspired by the CNN architecture, the QCNN is composed of a sequence of quantum convolutional and pooling layers. Each pooling layer typically reduces the number of qubits by a factor of two, thereby increasing the quantum circuit depth to <inline-formula id="inf1">
<mml:math id="m1">
<mml:mrow>
<mml:mi>O</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>log</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> for <inline-formula id="inf2">
<mml:math id="m2">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> input qubits. This logarithmic depth enables the implementation of extremely compact quantum machine learning models, with the number of parameters growing logarithmically with <inline-formula id="inf3">
<mml:math id="m3">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> [<xref ref-type="bibr" rid="B32">32</xref>, <xref ref-type="bibr" rid="B33">33</xref>, <xref ref-type="bibr" rid="B35">35</xref>]. Furthermore, QCNNs exhibit strong generalization capabilities [<xref ref-type="bibr" rid="B38">38</xref>] and are closely connected to tensor networks [<xref ref-type="bibr" rid="B29">29</xref>], making them an important architecture in QML.</p>
<p>The logarithmic circuit depth is one of the features that renders the QCNN an attractive architecture for NISQ devices, implying that the most natural design approach is to set the number of input qubits to a power of two. However, the number of input qubits required is determined by the input data dimension, i.e., the number of features in the data. If the input data require a number of qubits that is not a power of two, some layers will inevitably have odd numbers of qubits. This can occur either in the initial number of input qubits or during the pooling operation, representing a deviation from the optimal design and requiring appropriate adjustments. In particular, having an odd number of qubits in a quantum convolutional layer results in an increase in the circuit depth if all nearest-neighbor qubits interact with each other. Consequently, the run time increases and noise can negatively impact the overall performance and reliability of the QCNN. Moreover, it is unclear how breaking translational invariance in the pooling layer, a key property of the QCNN, affects overall performance. Because these considerations constrain the applicability of the QCNN algorithm, our goal is to optimize the QCNN architecture, developing an effective QML algorithm capable of handling arbitrary data dimensions.</p>
<p>In this study, we propose an efficient QCNN architecture capable of handling arbitrary data dimensions. Two naive approaches served as baselines to benchmark the proposed architectures: the classical data padding method, which increases the input data dimension through zero padding or periodic padding to encode it as a power of two, and the skip pooling method, which directly passes one qubit from each layer containing an odd number of qubits to the next layer without pooling. The first method requires additional ancillary qubits without increasing the circuit depth, whereas the second method does not require ancillary qubits but results in an increased circuit depth to preserve the translational invariance in the convolutional layers. By contrast, our proposed method effectively optimizes the QCNN architecture by applying a qubit padding technique that leverages ancillary qubits. By introducing an ancillary qubit into layers with an odd number of qubits, we can effectively construct convolutional layers without an additional increase in circuit depth. This enables a reduction in the total number of qubits by up to <inline-formula id="inf4">
<mml:math id="m4">
<mml:mrow>
<mml:mi>log</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> compared with the classical data padding method. Moreover, the reuse of a single ancillary qubit across multiple layers further reduces the required number of ancillary qubits. This strategy of qubit reuse efficiently optimizes the number of ancillary qubits along with the circuit depth. In addition, recycling the ancilla qubit facilitates uniform operations across layers, systematically enhancing the stability and efficiency of the QCNN architecture. To validate our approach, we benchmarked our proposed method against naive methods on various datasets: MNIST, Landsat satellite, Fashion-MNIST and Ionosphere datasets. Numerical simulation results show that our proposed method achieves a high classification accuracy comparable to that of naive methods. Notably, our method significantly reduces the number of qubits used compared with classical data padding methods, providing substantial advantages in terms of resource efficiency. We also conducted noise simulations using information from an IBM quantum device that mimics the operations and characteristics of real quantum hardware. The noise simulation results demonstrate that the proposed method exhibits less performance degradation and lower variability under realistic noise conditions than the skip pooling method. This is a consequence of the skip pooling method requiring a larger circuit depth. Because the proposed method not only improves the runtime but also enhances robustness against noise, it serves as a fundamental building block for the effective applicability of QCNNs to real-world data with an arbitrary number of features.</p>
<p>The remainder of this paper is organized as follows. We introduce the foundational concepts of QML in <xref ref-type="sec" rid="s2">Section 2</xref>, focusing on principles underlying quantum neural networks (QNNs) and QCNNs. <xref ref-type="sec" rid="s3">Section 3</xref> presents the detailed design of a QCNN architecture capable of handling arbitrary data dimensions, including a comparative analysis between naive methods and our proposed methods. Simulation results are presented in <xref ref-type="sec" rid="s4">Section 4</xref> along with a comparative performance analysis of the naive and proposed methods under both noiseless and noisy conditions. <xref ref-type="sec" rid="s5">Section 5</xref> explores possible extensions of multi-qubit quantum convolutional operations. Finally, concluding remarks are presented in <xref ref-type="sec" rid="s6">Section 6</xref>.</p>
</sec>
<sec id="s2">
<title>2 Background</title>
<sec id="s2-1">
<title>2.1 Quantum neural network</title>
<p>A DNN is a machine learning model constructed by deeply stacking layers of neurons [<xref ref-type="bibr" rid="B39">39</xref>]. Using nonlinear activation functions&#x2014;such as the sigmoid, ReLU, and hyperbolic tangent functions&#x2014;the DNN can learn patterns in complex data to solve various problems with high performance. Although the mathematical foundation for the success of DNNs remains an active area of research [<xref ref-type="bibr" rid="B40">40</xref>], several studies, as well as the universal approximation theorem, have demonstrated that neural networks can approximate complex functions with arbitrary accuracy [<xref ref-type="bibr" rid="B41">41</xref>]. On the other hand, a QNN is a quantum machine learning model, where the data is propagated through a PQC in the form of a quantum state. The data can be either intrinsically quantum, if the data source is a quantum system, or classical. In the latter case, which is the primary focus of this work, the classical data first has to be mapped to a quantum state. Note that nonlinear transformation of the input data can occur during this data mapping step. Since the parameters of the PQC are real-valued and its output is differentiable with respect to the parameters, they are typically trained through classical optimizers, similar to how DNNs are trained. In this sense, the QNN-based ML is also known to be a quantum-classical hybrid approach. Quantum-classical hybrid approaches using PQCs are effective at shallow circuit depths [<xref ref-type="bibr" rid="B18">18</xref>], which significantly enhances their applicability to NISQ devices with limited numbers of qubits. In addition, the PQC can approximate a broad family of functions with arbitrary accuracy, making it a good machine learning model [<xref ref-type="bibr" rid="B16">16</xref>, <xref ref-type="bibr" rid="B21">21</xref>].</p>
<p>A QNN consists of three primary components: (1) Data Embedding, (2) Data Processing, and (3)Measurements. These models transform classical data into quantum states, to be processed using a sequence of parameterized quantum gates. The training process connects the measurement results to the loss function, which is used to tune and train the parameters. <xref ref-type="fig" rid="F1">Figure 1</xref> depicts the overall training process of a QNN. Consider a dataset <inline-formula id="inf5">
<mml:math id="m5">
<mml:mrow>
<mml:mi mathvariant="script">D</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> with <inline-formula id="inf6">
<mml:math id="m6">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mi mathvariant="double-struck">R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula id="inf7">
<mml:math id="m7">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:mi mathvariant="double-struck">R</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, and PQC <inline-formula id="inf8">
<mml:math id="m8">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>, where <inline-formula id="inf9">
<mml:math id="m9">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
<mml:mo>,</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> represents a set of tunable parameters. Typically, the dataset is embedded into a quantum Hilbert space by a unitary transformation applied to <inline-formula id="inf10">
<mml:math id="m10">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> qubits initially prepared in <inline-formula id="inf11">
<mml:math id="m11">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mn>0</mml:mn>
<mml:mo stretchy="false">&#x232a;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2297;</mml:mo>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>. Denoting the data embedded state as <inline-formula id="inf12">
<mml:math id="m12">
<mml:mrow>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3c8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mtext>in</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">&#x232a;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>, the final state of the QNN can be expressed as follows:<disp-formula id="e1">
<mml:math id="m13">
<mml:mrow>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3c8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mtext>out</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">&#x232a;</mml:mo>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>l</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>l</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x2026;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3c8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mtext>in</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">&#x232a;</mml:mo>
<mml:mo>.</mml:mo>
</mml:mrow>
</mml:math>
</disp-formula>
</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>A schematic of the training process of a QNN. This figure outlines the sequential steps involved in training a QNN, starting with data preparation and qubit initialization, followed by the application of quantum gates to embed the data, and using a PQC for training. The output is obtained from measurements of the quantum state. Parameters are updated by minimizing the loss function.</p>
</caption>
<graphic xlink:href="fphy-13-1529188-g001.tif"/>
</fig>
<p>The output function of the QNN is <inline-formula id="inf13">
<mml:math id="m14">
<mml:mrow>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3b8;</mml:mi>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mo stretchy="false">&#x27e8;</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3c8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mtext>out</mml:mtext>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mi mathvariant="script">O</mml:mi>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3c8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mtext>out</mml:mtext>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">&#x27e9;</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>, where <inline-formula id="inf14">
<mml:math id="m15">
<mml:mrow>
<mml:mi mathvariant="script">O</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is an observable of the quantum circuit. The parameters are optimized using classical methods such as the gradient descent algorithm [<xref ref-type="bibr" rid="B42">42</xref>], which minimizes the following loss function:<disp-formula id="e2">
<mml:math id="m16">
<mml:mrow>
<mml:mi>L</mml:mi>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3b8;</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:munderover>
</mml:mstyle>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi mathvariant="bold-italic">&#x3b8;</mml:mi>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="bold-italic">x</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="false">&#x7c;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msup>
<mml:mo>.</mml:mo>
</mml:mrow>
</mml:math>
</disp-formula>
</p>
<p>Furthermore, gradients in quantum computing can be computed directly using methods such as parameter-shift rules, wherein derivatives are approximated by shifting the parameters at fixed intervals and then measuring the difference in the output of the quantum circuit as a result of that change [<xref ref-type="bibr" rid="B43">43</xref>, <xref ref-type="bibr" rid="B44">44</xref>]. However, as the number of qubits in the training PQC increases, the parameter space of the quantum circuit also increases, leading to the BP phenomenon [<xref ref-type="bibr" rid="B26">26</xref>]. Although this phenomenon represents a significant performance limitation of QML, it can be mitigated by applying hierarchical structures in the quantum circuits [<xref ref-type="bibr" rid="B29">29</xref>], such as the QCNN [<xref ref-type="bibr" rid="B30">30</xref>].</p>
</sec>
<sec id="s2-2">
<title>2.2 Quantum convolutional neural network</title>
<p>The QCNN is a type of PQC inspired by the concept of CNNs. QCNNs exhibit the property of translational invariance, with quantum circuits sharing the same parameters within the convolutional layer, and reduce dimensionality by tracing out some qubits during the pooling operation. A primary distinction between a QCNN and a CNN is that data in a QCNN are defined in a Hilbert space that grows exponentially with the number of qubits. Consequently, whereas classical convolution operations typically transform vectors into scalars, quantum convolution operations perform more complex linear mapping, transforming vectors into vectors through a unitary transformation of the state vector. Thus, quantum convolutional operations are distinct from classical convolutional operations. Problems defined in the exponentially large Hilbert space are intractable in a classical setting; however, QCNNs offer the possibility of effectively overcoming these challenges by utilizing qubits in a quantum setting. QCNNs have also demonstrated the capability to classify images in a manner similar to their classical counterparts [<xref ref-type="bibr" rid="B32">32</xref>]. In the convolutional layer, local features are extracted through unitary single-qubit rotations and entanglements between adjacent qubits, and the features dimensions are reduced in a pooling layer. Typically, the pooling layer includes parameterized two-qubit controlled unitary gates, and the control qubit is traced out after the gate operation to halve it. In binary classification tasks (i.e., <inline-formula id="inf15">
<mml:math id="m17">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:mrow>
<mml:mo stretchy="false">{</mml:mo>
<mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>,</mml:mo>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>), a QCNN repeats the convolutional and pooling layers until only one qubit remains, and performs classification by measuring the last qubit. These architectures maintain a shallow circuit depth by effectively reducing the number of qubits through a hierarchical structure, which is crucial for improving model performance and avoiding BP. In addition, by turning off the translational invariance property, which involves sharing the same parameters within the convolutional or pooling layer, more parameters can be introduced to the QML model while preserving the absence of BP. The structure between layers ensures a circuit depth of <inline-formula id="inf16">
<mml:math id="m18">
<mml:mrow>
<mml:mi>O</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>log</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> for <inline-formula id="inf17">
<mml:math id="m19">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> input qubits. In particular, the shallow depth of the QCNN contributes to its high performance in NISQ devices. In addition, the simple structure of repetitive circuits allows for the applicability of QCNNs to a wide range of tasks including classical and quantum data classification [<xref ref-type="bibr" rid="B31">31</xref>, <xref ref-type="bibr" rid="B32">32</xref>, <xref ref-type="bibr" rid="B45">45</xref>], error correction [<xref ref-type="bibr" rid="B31">31</xref>], and classical to quantum transfer learning [<xref ref-type="bibr" rid="B33">33</xref>, <xref ref-type="bibr" rid="B46">46</xref>]. Moreover, the QCNN architecture can be easily integrated into other QNN models and tasks, such as quantum recurrent neural networks [<xref ref-type="bibr" rid="B47">47</xref>], quantum generative adversarial networks [<xref ref-type="bibr" rid="B48">48</xref>, <xref ref-type="bibr" rid="B49">49</xref>], quantum graph convolutional neural networks [<xref ref-type="bibr" rid="B50">50</xref>, <xref ref-type="bibr" rid="B51">51</xref>], quantum one-class classifiers [<xref ref-type="bibr" rid="B52">52</xref>], and quantum self-supervised learning [<xref ref-type="bibr" rid="B53">53</xref>], bringing the aforementioned advantages to these areas as well.</p>
<p>It is also worth noting that mid-circuit measurements are not strictly required for QCNNs to achieve their key features&#x2014;translational invariance and dimensionality reduction. These characteristics can be effectively realized through specific arrangements of two-qubit gates, along with local measurements that commute with partial trace operations.</p>
</sec>
</sec>
<sec id="s3">
<title>3 QCNN architectures for arbitrary data dimensions</title>
<p>Data encoding is a crucial process in quantum computing that transforms classical data into quantum-state. During this process, the input data dimensions determine the number of qubits required to represent the quantum state. For example, amplitude encoding allows the input data <inline-formula id="inf18">
<mml:math id="m20">
<mml:mrow>
<mml:mi mathvariant="bold-italic">x</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>x</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>x</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mi>T</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>&#x2208;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mi mathvariant="double-struck">R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> with dimensions <inline-formula id="inf19">
<mml:math id="m21">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> to be represented as amplitudes of an <inline-formula id="inf20">
<mml:math id="m22">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>-qubit quantum state. Since a pooling layer in QCNN discards half of the qubits, <inline-formula id="inf21">
<mml:math id="m23">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> also has to be a power of two. The condition requiring the number of input qubits to be a power of two plays a critical role in the efficient applicability of QCNN algorithms. However, the dimensions of classical data do not always conform to this condition. In this section, we describe our proposed method that enables QCNN architectures to handle arbitrary data dimensions, along with naive baseline methods. To demonstrate the efficiency of the proposed method, we analyzed the number of ancillary qubits, parameters, and circuit depths when implementing the QCNN algorithm.</p>
<sec id="s3-1">
<title>3.1 Naive methods</title>
<sec id="s3-1-1">
<title>3.1.1 Classical data padding</title>
<p>In a CNN, the direct application of kernels to input feature maps during convolution operations can reduce the output feature map size compared to that of the input, leading to the potential loss of important information. Padding techniques that artificially enlarge the input feature map by adding specific values (typically zeros) around the input data are employed to prevent information loss and enhance model training. Inspired by CNN, similar padding strategies can be applied in QCNNs by increasing the data dimensions until the data can be encoded into a number of input qubits with a power of two, by either adding a constant value (zero) or periodically repeating the input data. These methods are referred to as &#x2018;zero-data padding&#x2019; and &#x2018;periodic-data padding&#x2019;, respectively. <xref ref-type="fig" rid="F2">Figure 2</xref> illustrates an example of the classical data padding method, with a handwritten digit image reduced to 30 dimensions. This 30-dimensional input can be encoded into five input qubits using amplitude encoding. Classical data padding methods can be applied to expand the data dimension, aligning the number of input qubits to a power of two. By using either zero- or periodic-data padding, the data dimensions can be expanded to <inline-formula id="inf22">
<mml:math id="m24">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>8</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>. Through these padding procedures, three additional ancillary qubits are introduced to the original system of five input qubits, resulting in an eight-qubit QCNN structure. Classical data padding approaches not only increase the number of qubits but may also result in poor classification performance because the number of dummy features that are added can often be significantly larger than the number of data features.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Schematic of the QCNN algorithm with eight input qubits using the classical data padding method. The quantum circuit consists of three components: data embedding (green squares), convolutional gate (blue squares), and pooling gate (red circles). The data embedding component is further divided into two methods: top embedding, which pads the zero-data, and bottom embedding, which pads the input data repeatedly. The convolutional and pooling gate use a PQC. Throughout the hierarchy, the convolutional gate consistently applies the same two-qubit ansatz to the nearest-neighboring qubit in each layer. In the <inline-formula id="inf23">
<mml:math id="m25">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>th convolutional layer, the set of gates that completes a loop connecting all nearest-neighboring qubits and the qubits at the boundaries can be repeated <inline-formula id="inf24">
<mml:math id="m26">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> times. The pooling gate uses the same approach, and can be represented as a controlled unitary transformation that is activated when the control qubit is 1.</p>
</caption>
<graphic xlink:href="fphy-13-1529188-g002.tif"/>
</fig>
<p>We denote the initial number of input qubits as <inline-formula id="inf25">
<mml:math id="m27">
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, and let <inline-formula id="inf26">
<mml:math id="m28">
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>log</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>m</mml:mi>
</mml:math>
</inline-formula>. Classical data padding typically employs <inline-formula id="inf27">
<mml:math id="m29">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>K</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> ancillary qubits. Without loss of generality, we assume that both the convolutional and pooling gates of the QCNN have one parameter and a depth of 1. The quantum circuit depth is <inline-formula id="inf28">
<mml:math id="m30">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>, where <inline-formula id="inf29">
<mml:math id="m31">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> denotes the number of complete sets of two-qubit gates connecting the nearest-neighboring qubits as well as the top and bottom qubits in the <inline-formula id="inf30">
<mml:math id="m32">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>th convolutional layer as depicted in <xref ref-type="fig" rid="F2">Figure 2</xref>. If the parameters are shared, then the total number of parameters is <inline-formula id="inf31">
<mml:math id="m33">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>. However, if the parameters are not shared, then the total number of parameters is <inline-formula id="inf32">
<mml:math id="m34">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
</sec>
<sec id="s3-1-2">
<title>3.1.2 Skip pooling</title>
<p>An alternate method enables implementation of the QCNN algorithm without the use of additional ancillary qubits. This method adopts a strategy where, within the QCNN structure, a qubit from each layer containing an odd number of qubits is passed directly to the next layer without performing a pooling operation. Although this approach minimizes the use of qubits, it inherently increases the circuit depth during convolutional operations in layers with odd numbers of qubits. This can affect the overall efficiency and execution speed of the quantum circuits. We refer to this method as &#x2018;skip pooling&#x2019;. <xref ref-type="fig" rid="F3">Figure 3A</xref> depicts an example of skip pooling, with a 30-dimensional input encoded into five input qubits using amplitude encoding. In the first layer of the QCNN, the convolutional operation between neighboring qubits introduces one more gate over classical data padding. Then, the <inline-formula id="inf33">
<mml:math id="m35">
<mml:mrow>
<mml:mn>5</mml:mn>
<mml:mi mathvariant="normal">t</mml:mi>
<mml:mi mathvariant="normal">h</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> qubit passes directly to the next layer without any pooling operations. This procedure is repeated for the second layer. When applied to layers that contain an odd number of qubits, skip pooling increases circuit depth during the convolutional operation. The increased circuit depth can be affected by noise, which may potentially propagate through each layer, reducing the accuracy and reliability of the information. This directly affects the efficiency and performance of the QCNN, necessitating an effective optimization strategy.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Schematic of a QCNN algorithm with five initial input qubits in a circuit with and without ancillary qubits. <bold>(A)</bold> Uses a method called skip pooling to perform convolution and pooling operations between each qubit without ancillary qubits. <bold>(B)</bold> Uses two ancillary qubits to construct the QCNN in a method called layer-wise qubit padding. The first layer has five qubits. Because this is an odd number of layers, one ancillary qubit is used to perform convolution and pooling operations. The second layer has three qubits, and another ancillary qubit is used. <bold>(C)</bold> Uses only one ancillary qubit to construct the QCNN using a method called single-ancilla qubit padding. Unlike <bold>(B)</bold>, the single ancillary qubit performs the convolution and pooling operations sequentially.</p>
</caption>
<graphic xlink:href="fphy-13-1529188-g003.tif"/>
</fig>
<p>Unlike classical data padding, no ancillary qubits are used; however, an additional circuit depth of <inline-formula id="inf34">
<mml:math id="m36">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is incurred, where <inline-formula id="inf35">
<mml:math id="m37">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2254;</mml:mo>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mi>mod</mml:mi>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> is an odd number of qubits in the <inline-formula id="inf36">
<mml:math id="m38">
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
<mml:mi mathvariant="normal">h</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> layer divided by the power of two. If the parameters are shared, then the total number of parameters is <inline-formula id="inf37">
<mml:math id="m39">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>. In contrast, if the parameters are not shared, then the total number of parameters is <inline-formula id="inf38">
<mml:math id="m40">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
</sec>
</sec>
<sec id="s3-2">
<title>3.2 Proposed methods</title>
<sec id="s3-2-1">
<title>3.2.1 Layer-wise qubit padding</title>
<p>As an alternative to the aforementioned naive methods, we introduce a qubit padding method that leverages ancillary qubits in the QCNN algorithm. Whereas classical data padding requires additional ancillary qubits by increasing the size of the input data, qubit padding directly leverages ancillary qubits in the convolutional and pooling operations of the QCNN. Using ancillary qubits for layers containing odd numbers of qubits, we optimized the QCNN algorithm and designed an architecture capable of handling arbitrary data dimensions. We refer to this method as &#x2018;layer-wise qubit padding&#x2019;. <xref ref-type="fig" rid="F3">Figure 3B</xref> depicts an example of layer-wise qubit padding. As in the skip pooling example, a 30-dimensional input was encoded into five input qubits using amplitude encoding. In the first layer of the QCNN, one ancillary qubit is added to perform convolutional and pooling operations with neighboring qubits. This ensures pairwise matching between all qubits in two steps, avoiding additional circuit depth that may arise in skip pooling. This procedure is repeated for the second layer. Consequently, in a five-qubit QCNN with layer-wise qubit padding and two layers containing an odd number of qubits, two ancillary qubits are used to optimize the architecture.</p>
<p>Layer-wise qubit padding generally requires <inline-formula id="inf39">
<mml:math id="m41">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> ancillary qubits. The quantum circuit depth of <inline-formula id="inf40">
<mml:math id="m42">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is identical to that used in classical data padding. If the parameters are shared, then the total number of parameters is <inline-formula id="inf41">
<mml:math id="m43">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>. On the other hand, if the parameters are not shared, then the total number of parameters is <inline-formula id="inf42">
<mml:math id="m44">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>, which is <inline-formula id="inf43">
<mml:math id="m45">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> more than that for skip pooling.</p>
</sec>
<sec id="s3-2-2">
<title>3.2.2 Single-ancilla qubit padding</title>
<p>Finally, we propose &#x2018;single-ancilla qubit padding,&#x2019; a QCNN architecture designed to handle arbitrary input data dimensions using only one ancillary qubit. By reusing the ancillary qubit throughout the QCNN architecture, we significantly reduced the number of total qubits required for optimization. <xref ref-type="fig" rid="F3">Figure 3C</xref> illustrates an example of single-ancilla qubit padding. Unlike layer-wise qubit padding, this method reuses the same ancillary qubit for every layer with an odd number of qubits. Preserving the information of the ancillary qubit without resetting it when it is passed to the next layer plays a crucial role in enhancing the stability and performance of model training.</p>
<p>Although single-ancilla qubit padding uses only one ancillary qubit, the quantum circuit depth and number of parameters remain the same as those in layer-wise qubit padding. <xref ref-type="fig" rid="F4">Figure 4A</xref> illustrates the circuit depths of the skip pooling and qubit padding methods. As the number of input qubits increases, the circuit depth of the qubit padding method is logarithmically less than that of the skip pooling method. <xref ref-type="fig" rid="F4">Figure 4B</xref> illustrates the number of parameters in the case of parameter-sharing off for the classical data padding, skip pooling, and layer-wise and single-ancilla qubit padding methods. In the case of parameter-sharing on, the number of parameters is the same across all methods, hence we do not track how the number of parameters changes. We only compared the number of parameters in the case of parameter-sharing off. Without loss of generality, we assumed the convolutional layer <inline-formula id="inf44">
<mml:math id="m46">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> to be equal to 1. When the number of input qubits is not a power of two, classical data padding uses the largest number of the parameters, whereas skip pooling and qubit padding use relatively fewer parameters. By introducing additional qubits into certain layers, the qubit padding method uses slightly more parameters than skip pooling. Therefore, single-ancilla qubit padding enables the design of efficient QCNN architectures with optimal allocation of quantum resources such as ancillary qubits and quantum gates.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>
<bold>(A)</bold> Semi-log plot illustrating the difference in circuit depth between the skip pooling method and layer-wise and single-ancilla qubit padding method. The dashed line represents circuit depth in the skip pooling method, the dash-dot line denotes circuit depth in the layer-wise and single-ancilla qubit padding method, the solid line represents the difference between the two methods, and the dotted line corresponds to <inline-formula id="inf45">
<mml:math id="m47">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>log</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, provided as a guide to the eye. <bold>(B)</bold> Semi-log plot illustrating the number of parameters in the case of parameter sharing off for the classical data padding, skip pooling, and layer-wise and single-ancilla qubit padding methods. The solid line denotes the number of parameters for classical data padding, the dashed line represents the number of parameters for the skip pooling method, and the dash-dot line corresponds to the number of parameters for the layer-wise and single-ancilla qubit padding.</p>
</caption>
<graphic xlink:href="fphy-13-1529188-g004.tif"/>
</fig>
</sec>
</sec>
</sec>
<sec sec-type="results" id="s4">
<title>4 Results</title>
<p>The previous section provided an overview of various QCNN padding methods. In this section, we benchmark and evaluate our proposed padding methods in comparison with naive methods using a variety of classical datasets. To address the characteristics of NISQ devices, we added noise to the quantum circuits for benchmarking.</p>
<sec id="s4-1">
<title>4.1 Methods and setup</title>
<sec id="s4-1-1">
<title>4.1.1 Datasets</title>
<p>Our experiments were conducted using a variety of datasets. The MNIST dataset consists of handwritten digits, each represented as a <inline-formula id="inf46">
<mml:math id="m48">
<mml:mrow>
<mml:mn>28</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>28</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> pixel image in grayscale [<xref ref-type="bibr" rid="B54">54</xref>]. The dataset consists of a total of 60,000 training and 10,000 test images, each labeled with a numerical value ranging from 0 to 9. In our benchmarking, we focused on binary classification tasks by selecting two distinct pairs of labels: 0 &#x26; 1 and 5 &#x26; 6. In the noiseless scenario, we split the dataset into 10,000 training, 1,000 validation, and 1,000 test sets. In the noisy scenario, we split the dataset into 2,000 training, 200 validation, and 200 test sets. Furthermore, <inline-formula id="inf47">
<mml:math id="m49">
<mml:mrow>
<mml:mn>28</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>28</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> features were relatively high-dimensional for current quantum hardware; therefore, we used principal component analysis as a dimensionality reduction technique to reduce the dataset to 30 features.</p>
<p>In addition, we conducted experiments with the Landsat Satellite, Fashion-MNIST and Ionosphere datasets in the noisy scenario. The Landsat Satellite dataset classifies multi-spectral values of pixels from satellite images [<xref ref-type="bibr" rid="B55">55</xref>]. The dataset contains a total of 6,435 instances with 36 features, where each instance is labeled into one of six land cover classes. In our benchmarking, we focused on binary classification tasks by selecting two distinct pairs of labels: 1 (Red Soil) and 2 (Cotton Crop). We split the dataset into 1,000 training, 124 validation, and 200 test instances, then reduced the 36 features to 30 features using principal component analysis.</p>
<p>The Fashion-MNIST dataset consists of <inline-formula id="inf48">
<mml:math id="m50">
<mml:mrow>
<mml:mn>28</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>28</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> grayscale images of clothing [<xref ref-type="bibr" rid="B56">56</xref>]. The dataset consists of a total of 60,000 training and 10,000 test images, each labeled with one of 10 clothing classes. In our benchmarking, we focused on binary classification tasks by selecting two distinct pairs of labels: 0 &#x26; 1 and 0 &#x26; 2. We split the dataset into 2,000 training, 200 validation, and 200 test sets, then reduced the <inline-formula id="inf49">
<mml:math id="m51">
<mml:mrow>
<mml:mn>28</mml:mn>
<mml:mo>&#xd7;</mml:mo>
<mml:mn>28</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> features to 30 features using principal component analysis.</p>
<p>The Ionosphere dataset classifies radar returns from the ionosphere [<xref ref-type="bibr" rid="B57">57</xref>]. The dataset contains a total of 351 instances with 34 features labeled as &#x2018;good&#x2019; or &#x2018;bad&#x2019;, where &#x2018;good&#x2019; indicates radar returns that show evidence of structure in the ionosphere, and &#x2018;bad&#x2019; indicates signals that pass through without detecting any structure. We divided the data between 251 training and 51 test sets.</p>
</sec>
<sec id="s4-1-2">
<title>4.1.2 Ansatz</title>
<p>We tested two different structures of the parameterized quantum circuit, also referred to as the ansatz, for the convolutional operations. The first one consists of two parameterized single-qubit rotations and a CNOT gate, as shown in <xref ref-type="fig" rid="F5">Figure 5A</xref> [<xref ref-type="bibr" rid="B29">29</xref>]. This represents the simplest two-qubit ansatz. The second one is designed to express an arbitrary two-qubit unitary transformation. In general, any two-qubit unitary gate in the SU(4) group can be decomposed using at most three CNOT gates and 15 elementary single-qubit gates [<xref ref-type="bibr" rid="B45">45</xref>, <xref ref-type="bibr" rid="B58">58</xref>]. The quantum circuits shown in <xref ref-type="fig" rid="F5">Figure 5C</xref> represent the parameterization of an arbitrary SU(4) gate. <xref ref-type="fig" rid="F5">Figure 5B</xref> shows the pooling circuit, where two controlled rotations, <inline-formula id="inf50">
<mml:math id="m52">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>y</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula id="inf51">
<mml:math id="m53">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>y</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>, are applied, with each activated when the control qubit is 1 (filled circle) or 0 (open circle). <xref ref-type="sec" rid="s13">Supplementary Appendix SA</xref> presents definitions of the quantum gates used in this study. We constructed two ansatz sets using a different combination of the convolutional and pooling circuits. Ansatz set 1 uses convolutional circuit 1 and a pooling circuit as shown in the figure, whereas ansatz set 2 uses convolution circuit 2 and pooling without a parameterized circuit. In the latter case, the pooling performs the partial trace operations without parameterized gates because the convolution circuit one is expressive enough to implement any two-qubit unitary operation.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>PQC for convolution and pooling operations. <bold>(A)</bold> and <bold>(B)</bold> are the convolutional and pooling circuits, respectively, that compose ansatz set 1. <bold>(C)</bold> Is the convolutional circuit that composes ansatz set 2. <inline-formula id="inf52">
<mml:math id="m54">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is the rotation by <inline-formula id="inf53">
<mml:math id="m55">
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> around the <inline-formula id="inf54">
<mml:math id="m56">
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>-axis of the Bloch sphere, and <inline-formula id="inf55">
<mml:math id="m57">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>3</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>&#x3d5;</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>&#x3bb;</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is an arbitrary single-qubit gate, which can be expressed as <inline-formula id="inf56">
<mml:math id="m58">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>3</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>&#x3d5;</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>&#x3bb;</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>&#x3d5;</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>&#x3c0;</mml:mi>
<mml:mo>/</mml:mo>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>&#x3b8;</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>
<inline-formula id="inf156">
<mml:math id="m158">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>&#x3c0;</mml:mi>
<mml:mo>/</mml:mo>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>R</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>&#x3bb;</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
</caption>
<graphic xlink:href="fphy-13-1529188-g005.tif"/>
</fig>
</sec>
</sec>
<sec id="s4-2">
<title>4.2 Simulation without noise</title>
<p>In this section, we present numerical experimental results that evaluate the performance of QCNNs with various padding methods for binary classification tasks in a noiseless environment. <xref ref-type="table" rid="T1">Tables 1</xref>, <xref ref-type="table" rid="T2">2</xref> summarize the number of ancillary qubits, circuit depth, and number of parameters required for the naive and proposed methods. Classical data padding and skip pooling methods exhibit a significant difference in terms of the utilization of ancillary qubits. Specifically, classical data padding maximally uses ancillary qubits to apply a natural QCNN algorithm, whereas skip pooling does not use any ancillary qubits. However, skip pooling poses a potential drawback in the form of a potential increase in circuit depth, which affects computational power and runtime. Additionally, the use of fewer qubits results in the use of fewer parameters when they are not shared by each layer. However, the qubit padding method can apply an efficient QCNN algorithm with fewer qubits. Because the single-ancilla qubit padding method uses only one ancillary qubit, it does not incur additional circuit depth and offers the advantage of utilizing a slightly larger number of parameters than the skip-pooling method.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Comparison of ancillary qubits, circuit depth, and the total number of parameters for classical data padding and skip pooling. <inline-formula id="inf57">
<mml:math id="m59">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2254;</mml:mo>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mi>mod</mml:mi>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> is an odd number of qubits in the <inline-formula id="inf58">
<mml:math id="m60">
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
<mml:mi mathvariant="normal">h</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> layer when divided by a power of two.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Padding methods</th>
<th align="left"/>
<th align="center">Classical data padding</th>
<th align="center">Skip pooling</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">Ancillary qubits</td>
<td align="left"/>
<td align="center">
<inline-formula id="inf59">
<mml:math id="m61">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>K</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">0</td>
</tr>
<tr>
<td align="center">Circuit depth</td>
<td align="left"/>
<td align="center">
<inline-formula id="inf60">
<mml:math id="m62">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">
<inline-formula id="inf61">
<mml:math id="m63">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> &#x2b; <inline-formula id="inf62">
<mml:math id="m64">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
<tr>
<td rowspan="2" align="center">Parameters</td>
<td align="center">p-s on</td>
<td align="center">
<inline-formula id="inf63">
<mml:math id="m65">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">
<inline-formula id="inf64">
<mml:math id="m66">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
<tr>
<td align="center">p-s off</td>
<td align="center">
<inline-formula id="inf65">
<mml:math id="m67">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">
<inline-formula id="inf66">
<mml:math id="m68">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>The notation &#x201c;p-s on&#x201d; and &#x201c;p-s off&#x201d; indicates whether parameter-sharing is enabled or disabled, respectively.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Comparison of ancillary qubits, circuit depth, and the total number of parameters for layer-wise and single-ancilla qubit padding. <inline-formula id="inf67">
<mml:math id="m69">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2254;</mml:mo>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mi>mod</mml:mi>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> is an odd number of qubits in the <inline-formula id="inf68">
<mml:math id="m70">
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
<mml:mi mathvariant="normal">h</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> layer when divided by a power of two.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Padding methods</th>
<th align="left"/>
<th align="center">Layer-wise qubit padding</th>
<th align="center">Single-ancilla qubit padding</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">Ancillary qubits</td>
<td align="left"/>
<td align="center">
<inline-formula id="inf69">
<mml:math id="m71">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">1</td>
</tr>
<tr>
<td align="center">Circuit depth</td>
<td align="left"/>
<td align="center">
<inline-formula id="inf70">
<mml:math id="m72">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">
<inline-formula id="inf71">
<mml:math id="m73">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
<tr>
<td rowspan="2" align="center">Parameters</td>
<td align="center">p-s on</td>
<td align="center">
<inline-formula id="inf72">
<mml:math id="m74">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">
<inline-formula id="inf73">
<mml:math id="m75">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
<tr>
<td align="center">p-s off</td>
<td align="center">
<inline-formula id="inf74">
<mml:math id="m76">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="center">
<inline-formula id="inf75">
<mml:math id="m77">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x2308;</mml:mo>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>&#x2309;</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>Y</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>The notation &#x201c;p-s on&#x201d; and &#x201c;p-s off&#x201d; indicates whether parameter-sharing is enabled or disabled, respectively.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>The simulation results were based on experiments using two different ansatz sets, denoted as ansatz set 1 and ansatz set 2. <xref ref-type="table" rid="T3">Table 3</xref> lists the number of ancillary qubits, circuit depth, and parameters when applying different ansatz sets to the QCNN. All convolutional layers were used only once, and without loss of generality, the circuit depth was obtained by setting the convolution and pooling gate depths to 1. We obtained results from 10 repeated experiments with randomly initialized parameters for each ansatz set. The performance of the QCNN model was evaluated using the mean squared error (MSE) loss function. Model parameters were updated using the Adam optimizer [<xref ref-type="bibr" rid="B59">59</xref>]. The learning rate and batch size were set to 0.01 and 25, respectively. Training was performed with 10 epochs on the MNIST datasets.</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Comparison of the number of ancillary qubits, circuit depth, and total number of parameters for each padding method with different ansatz sets based on five initial input qubits. The notation &#x2018;p-s on&#x2019; and &#x2018;p-s off&#x2019; indicates whether parameter-sharing is enabled or disabled, respectively.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="2" align="center">Padding methods</th>
<th rowspan="2" align="center">Ancillary qubits</th>
<th rowspan="2" align="center">Circuit depth</th>
<th colspan="2" align="center">Parameters (p-s on)</th>
<th colspan="2" align="center">Parameters (p-s off)</th>
</tr>
<tr>
<th align="center">ansatz set 1</th>
<th align="center">ansatz set 2</th>
<th align="center">ansatz set 1</th>
<th align="center">ansatz set 2</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">Classical data padding</td>
<td align="center">3</td>
<td align="center">8</td>
<td align="center">12</td>
<td align="center">45</td>
<td align="center">40</td>
<td align="center">195</td>
</tr>
<tr>
<td align="center">Skip pooling</td>
<td align="center">0</td>
<td align="center">10</td>
<td align="center">12</td>
<td align="center">45</td>
<td align="center">26</td>
<td align="center">135</td>
</tr>
<tr>
<td align="center">Layer-wise qubit padding</td>
<td align="center">2</td>
<td align="center">8</td>
<td align="center">12</td>
<td align="center">45</td>
<td align="center">34</td>
<td align="center">165</td>
</tr>
<tr>
<td align="center">Single-ancilla qubit padding</td>
<td align="center">1</td>
<td align="center">8</td>
<td align="center">12</td>
<td align="center">45</td>
<td align="center">34</td>
<td align="center">165</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>The average accuracies and standard deviations of the binary classification task on the MNIST (labels 0 &#x26; 1 and 5 &#x26; 6) datasets obtained using ansatz set 1 are shown in <xref ref-type="fig" rid="F6">Figure 6</xref>. Both single-ancilla qubit padding and skip pooling achieved superior accuracy across the two datasets. For example, for labels 0 &#x26; 1 in the MNIST dataset, single-ancilla qubit padding with shared parameters achieved an average accuracy of 91.75 (<inline-formula id="inf76">
<mml:math id="m78">
<mml:mrow>
<mml:mo>&#xb1;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>2.57), whereas skip pooling showed an accuracy of 91.76 (<inline-formula id="inf77">
<mml:math id="m79">
<mml:mrow>
<mml:mo>&#xb1;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>2.98). Such a trend was consistently observed in all other test cases, indicating the effectiveness of these methods in handling classification tasks with greater precision. However, the other padding methods, although still effective, did not attain the same levels of accuracy as single-ancilla qubit padding and skip pooling. The result obtained using ansatz set 2 is shown in <xref ref-type="fig" rid="F7">Figure 7</xref>. Single-ancilla qubit padding and skip pooling consistently achieved higher performance. For example, for labels 5 &#x26; 6 in the MNIST dataset, single-ancilla qubit padding achieved an average accuracy of 93.07 (<inline-formula id="inf78">
<mml:math id="m80">
<mml:mrow>
<mml:mo>&#xb1;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>1.14) and skip pooling achieved 94.59 (<inline-formula id="inf79">
<mml:math id="m81">
<mml:mrow>
<mml:mo>&#xb1;</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>0.62), showing better results with fewer qubits. Although skip pooling is efficient in terms of qubit usage and performance, it results in a deeper circuit that may be more susceptible to noise, particularly in existing noisy quantum devices. In contrast, single-ancilla qubit padding is more robust to noise, potentially making it more suitable for implementation on quantum devices. This will be demonstrated in the following section.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>QCNN model performance with various padding methods constructed using ansatz set 1. The bar chart shows the average accuracy and standard deviation for <bold>(A)</bold> the MNIST 0 &#x26; 1 dataset, <bold>(B)</bold> and the MNIST 5 &#x26; 6 dataset. The x-axis differentiates between the case of parameter-sharing on and parameter-sharing off. Unfilled bars represent zero-data padding, forward slash bars represents periodic-data padding, backslash bars represents skip Pooling, horizontal dash bars represents layer-wise ancilla, and dots bars represent single-ancilla.</p>
</caption>
<graphic xlink:href="fphy-13-1529188-g006.tif"/>
</fig>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>QCNN model performance with various padding methods constructed using ansatz set 2. The bar chart shows the average accuracy and standard deviation for <bold>(A)</bold> the MNIST 0 &#x26; 1 dataset, and <bold>(B)</bold> the MNIST 5 &#x26; 6 dataset. The x-axis differentiates between the case of parameter-sharing on and parameter-sharing off. Unfilled bars represent zero-data padding, forward slash bars represents periodic-data padding, backslash bars represents skip pooling, horizontal dash bars represents layer-wise ancilla, and dots bars represent single-ancilla.</p>
</caption>
<graphic xlink:href="fphy-13-1529188-g007.tif"/>
</fig>
</sec>
<sec id="s4-3">
<title>4.3 Simulation with noise</title>
<p>We conducted an additional experiment to evaluate the impact of noise in a quantum computing environment on the performance of QCNN algorithms. In particular, we considered the influence of circuit depth on error accumulation in quantum computation by comparing performance between the single-ancilla qubit padding and skip pooling methods. The noise simulations focused on types of noise that closely relate to circuit depth, and state preparation and measurement errors (SPAM) were excluded, as they were beyond the scope of our interest in this study. We considered various types of noise in quantum devices, such as depolarization errors, gate lengths, and thermal relaxation, but not the physical connectivity of qubits. <xref ref-type="sec" rid="s13">Supplementary Appendix SC</xref> provides details of the noise circuits used in this experiment.</p>
<p>We used the noise parameters observed from IBMQ Jakarta, a real quantum device, to simulate a realistic noise model. <xref ref-type="table" rid="T4">Table 4</xref> lists the average error rates observed on IBMQ Jakarta. To evaluate the influence of a range of noise levels, we performed experiments with depolarizing errors and gate lengths ranging from one to five times the original values. This multiplication was applied consistently over 100 repeated experiments with randomly initialized parameters, and the results are shown in <xref ref-type="fig" rid="F8">Figure 8</xref>. For the MNIST, Landsat satellite, and Fashion-MNIST dataset, as noise levels increased, the single-ancilla qubit padding method showed less accuracy degradation compared to skip pooling. In the case of the Ionosphere dataset, the single-ancilla method consistently outperformed skip pooling across all tested noise levels. These results demonstrate that the proposed method provides an optimal solution for constructing an efficient QCNN architecture with minimal resource overhead, maintaining robust performance despite noise and imperfections.</p>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Average error rates for the IBM quantum device, <italic>ibmq_jakarta</italic>, utilized in the noisy simulation.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">1-Qubit depolarizing</th>
<th align="center">2-Qubit depolarizing</th>
<th align="center">1-Qubit gate length</th>
<th align="center">2-Qubit gate length</th>
<th align="center">
<inline-formula id="inf80">
<mml:math id="m82">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>T</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>
</th>
<th align="center">
<inline-formula id="inf81">
<mml:math id="m83">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>T</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>
</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">0.0004</td>
<td align="center">0.0126</td>
<td align="center">35.56 (ns)</td>
<td align="center">327.11 (ns)</td>
<td align="center">128.43 (us)</td>
<td align="center">33.85 (us)</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>Test results of QCNN model with skip pooling and single-ancilla qubit padding constructed with ansatz set 2. The depolarizing errors and gate lengths were increased from their original values (x1) up to a maximum of (x5). The solid blue line represents skip pooling, whereas the dashed red line denotes single-ancilla qubit padding. The mean and standard error were obtained from 100 repeated experiments with parameters initialized randomly. <bold>(A)</bold> Shows the average accuracy and standard error of classification for 0 &#x26; 1 in the MNIST dataset, <bold>(B)</bold> shows the average accuracy and standard error of classification for 5 &#x26; 6 in the MNIST dataset, and <bold>(C)</bold> shows the average accuracy and standard error of classification for 1 &#x26; 2 in the Landsat satellite dataset. <bold>(D)</bold> Shows the average accuracy and standard error of classification for 0 &#x26; 1 in the Fashion-MNIST dataset, <bold>(E)</bold> shows the average accuracy and standard error of classification for 0 &#x26; 2 in the Fashion-MNIST dataset. <bold>(F)</bold> Shows the average accuracy and standard error of classification for the Ionosphere dataset.</p>
</caption>
<graphic xlink:href="fphy-13-1529188-g008.tif"/>
</fig>
</sec>
</sec>
<sec id="s5">
<title>5 Extension to multi-qubit quantum convolutional operations</title>
<p>An arbitrary unitary operation acting on <inline-formula id="inf82">
<mml:math id="m84">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> qubits, which is an element in the <inline-formula id="inf83">
<mml:math id="m85">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>U</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> group, can be specified using <inline-formula id="inf84">
<mml:math id="m86">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mn>4</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> real parameters. This implies that the number of elementary gates required to implement an arbitrary <inline-formula id="inf85">
<mml:math id="m87">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>-qubit unitary operation increases exponentially with <inline-formula id="inf86">
<mml:math id="m88">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>. Therefore, minimizing the number of qubits involved in a quantum convolutional operation is beneficial in practice. This is a primary motivation for designing quantum convolutional operations that act on only two qubits, as considered in this study. However, quantum convolutional operations can theoretically act on any <inline-formula id="inf87">
<mml:math id="m89">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> number of qubits. In general, the quantum circuit depth of any given convolutional layer, denoted by <inline-formula id="inf88">
<mml:math id="m90">
<mml:mrow>
<mml:mi>l</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, is greater than or equal to <inline-formula id="inf89">
<mml:math id="m91">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> (i.e., <inline-formula id="inf90">
<mml:math id="m92">
<mml:mrow>
<mml:mi>l</mml:mi>
<mml:mo>&#x2265;</mml:mo>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>), and equality is satisfied only when the quantum convolutional layer consists of <inline-formula id="inf91">
<mml:math id="m93">
<mml:mrow>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> qubits where <inline-formula id="inf92">
<mml:math id="m94">
<mml:mrow>
<mml:mi>a</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi mathvariant="double-struck">Z</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2b;</mml:mo>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is a positive integer (i.e., <inline-formula id="inf93">
<mml:math id="m95">
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is a positive-integer multiple of <inline-formula id="inf94">
<mml:math id="m96">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>). Therefore, if the number of qubits in a quantum convolutional layer, denoted by <inline-formula id="inf95">
<mml:math id="m97">
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, is not an integer multiple of <inline-formula id="inf96">
<mml:math id="m98">
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, the circuit depth of the given convolutional layer can be minimized at the cost of introducing <inline-formula id="inf97">
<mml:math id="m99">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2032;</mml:mo>
</mml:mrow>
</mml:msup>
<mml:mo>&#x3c;</mml:mo>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> ancilla qubits such that <inline-formula id="inf98">
<mml:math id="m100">
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2032;</mml:mo>
</mml:mrow>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
<p>For example, consider a quantum convolutional layer consisting of <inline-formula id="inf99">
<mml:math id="m101">
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>7</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> qubits where each quantum convolutional operation acts on <inline-formula id="inf100">
<mml:math id="m102">
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>3</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> qubits. By utilizing <inline-formula id="inf101">
<mml:math id="m103">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mo>&#x2032;</mml:mo>
</mml:mrow>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula> ancilla qubits, the circuit depth of the quantum convolutional layer can be minimized to three.</p>
</sec>
<sec sec-type="conclusion" id="s6">
<title>6 Conclusion</title>
<p>In this study, we designed a QCNN architecture that can handle arbitrary data dimensions. Using qubit padding, we optimized the allocation of quantum resources through the efficient use of ancillary qubits. Our method not only reduces the number of ancillary qubits, but also optimizes the circuit depth to construct an efficient QCNN architecture. This results in an optimal solution that is computationally efficient and robust against noise. We benchmarked the performance of our QCNN using both naive methods and the proposed methods on various datasets for binary classification. In simulations without noise, both skip pooling and our proposed single-ancilla qubit padding method achieved high accuracy in most cases. We also compared performance between single-ancilla qubit padding and skip pooling in a noisy simulation, using the noise model and parameters of an IBM quantum device. Our results demonstrate that as the noise level increases, single-ancilla qubit padding exhibits less performance degradation and lower sensitivity to variation. Therefore, the proposed method serves as a fundamental building block for the effective application of QCNN to real-world data of arbitrary input dimension.</p>
<p>The main focus of our study is on the analysis of classical data using QML, reflecting the prevalence of classical datasets in modern society. Nevertheless, data can also be intrinsically quantum [<xref ref-type="bibr" rid="B60">60</xref>]. In such cases, classical dimensionality reduction or data padding to adjust the number of input qubits to a power of two is not feasible. However, the single-ancilla qubit padding method can be easily adapted and remain valuable.</p>
<p>As a final remark, our work aims to guide users in selecting the optimal QCNN circuit design with respect to their specific requirements and system environment. We do not intend to rule out the skip-pooling method; it remains a viable option if increasing the number of qubits is more challenging than increasing the circuit depth. Conversely, if minimizing circuit depth is critical and adding an extra qubit is relatively easy, then the single-ancilla method would be preferable. Additionally, if the task at hand requires a higher model complexity, the single-ancilla method with parameter-sharing off can be used.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s7">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/<xref ref-type="sec" rid="s13">Supplementary Material</xref>, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec sec-type="author-contributions" id="s8">
<title>Author contributions</title>
<p>CL: Formal Analysis, Investigation, Methodology, Software, Validation, Writing&#x2013;original draft, Writing&#x2013;review and editing. IA: Investigation, Software, Writing&#x2013;review and editing. DK: Investigation, Methodology, Software, Writing&#x2013;original draft. JL: Investigation, Software, Writing&#x2013;original draft. SP: Investigation, Methodology, Software, Writing&#x2013;original draft, Writing&#x2013;review and editing. J-YR: Investigation, Methodology, Software, Writing&#x2013;original draft. DP: Conceptualization, Formal Analysis, Funding acquisition, Methodology, Supervision, Writing&#x2013;original draft, Writing&#x2013;review and editing.</p>
</sec>
<sec sec-type="funding-information" id="s9">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research, authorship, and/or publication of this article. This research was supported by Institute for Information &#x26; Communications Technology Promotion (IITP) grant funded by the Korea government (No. 2019-0-00003, Research and Development of Core Technologies for Programming, Running, Implementing and Validating of Fault-Tolerant Quantum Computing System), the Yonsei University Research Fund of 2024 (2024-22-0147), the National Research Foundation of Korea (2022M3E4A1074591, 2023M3K5A1094813), the KIST Institutional Program (2E32941-24-008), and the Ministry of Trade, Industry, and Energy (MOTIE), Korea, under the Industrial Innovation Infrastructure Development Project (Project No. RS-2024-00466693).</p>
</sec>
<sec sec-type="COI-statement" id="s10">
<title>Conflict of interest</title>
<p>Author J-YR was employed by Norma Inc.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="s11">
<title>Generative AI statement</title>
<p>The author(s) declare that no Generative AI was used in the creation of this manuscript.</p>
</sec>
<sec sec-type="disclaimer" id="s12">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="s13">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fphy.2025.1529188/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fphy.2025.1529188/full&#x23;supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Supplementaryfile1.pdf" id="SM1" mimetype="application/pdf" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<label>1.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jain</surname>
<given-names>AK</given-names>
</name>
<name>
<surname>Mao</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Mohiuddin</surname>
<given-names>KM</given-names>
</name>
</person-group>. <article-title>Artificial neural networks: a tutorial</article-title>. <source>Computer</source> (<year>1996</year>) <volume>29</volume>(<issue>3</issue>):<fpage>31</fpage>&#x2013;<lpage>44</lpage>. <pub-id pub-id-type="doi">10.1109/2.485891</pub-id>
</citation>
</ref>
<ref id="B2">
<label>2.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vaswani</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Shazeer</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Parmar</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Uszkoreit</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Jones</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Gomez</surname>
<given-names>AN</given-names>
</name>
<etal/>
</person-group> <article-title>Attention is all you need</article-title>. <source>Adv Neural Inf Process Syst</source> (<year>2017</year>) <volume>30</volume>.</citation>
</ref>
<ref id="B3">
<label>3.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yann</surname>
<given-names>LC</given-names>
</name>
<name>
<surname>Bernhard</surname>
<given-names>B</given-names>
</name>
<name>
<surname>John</surname>
<given-names>SD</given-names>
</name>
<name>
<surname>Henderson</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Howard</surname>
<given-names>RE</given-names>
</name>
<name>
<surname>Hubbard</surname>
<given-names>W</given-names>
</name>
<etal/>
</person-group> <article-title>Backpropagation applied to handwritten zip code recognition</article-title>. <source>Neural Comput</source> (<year>1989</year>) <volume>1</volume>(<issue>4</issue>):<fpage>541</fpage>&#x2013;<lpage>51</lpage>. <pub-id pub-id-type="doi">10.1162/neco.1989.1.4.541</pub-id>
</citation>
</ref>
<ref id="B4">
<label>4.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>LeCun</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Bottou</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Bengio</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Haffner</surname>
<given-names>P</given-names>
</name>
</person-group>. <article-title>Gradient-based learning applied to document recognition</article-title>. <source>Proc IEEE</source> (<year>1998</year>) <volume>86</volume>(<issue>11</issue>):<fpage>2278</fpage>&#x2013;<lpage>324</lpage>. <pub-id pub-id-type="doi">10.1109/5.726791</pub-id>
</citation>
</ref>
<ref id="B5">
<label>5.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Krizhevsky</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Sutskever</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Hinton</surname>
<given-names>GE</given-names>
</name>
</person-group>. <article-title>Imagenet classification with deep convolutional neural networks</article-title>. <source>Adv Neural Inf Process Syst</source> (<year>2012</year>) <volume>25</volume>.</citation>
</ref>
<ref id="B6">
<label>6.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Ren</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Deep residual learning for image recognition</article-title>. In: <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source> (<year>2016</year>). p. <fpage>770</fpage>&#x2013;<lpage>8</lpage>.</citation>
</ref>
<ref id="B7">
<label>7.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Girshick</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Donahue</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Darrell</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Malik</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Rich feature hierarchies for accurate object detection and semantic segmentation</article-title>. In: <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source> (<year>2014</year>). p. <fpage>580</fpage>&#x2013;<lpage>7</lpage>.</citation>
</ref>
<ref id="B8">
<label>8.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Shen</surname>
<given-names>X</given-names>
</name>
<name>
<surname>Brandt</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Hua</surname>
<given-names>G</given-names>
</name>
</person-group>. <article-title>A convolutional neural network cascade for face detection</article-title>. In: <source>Proceedings of the IEEE conference on computer vision and pattern recognition</source> (<year>2015</year>) <fpage>5325</fpage>&#x2013;<lpage>34</lpage>.</citation>
</ref>
<ref id="B9">
<label>9.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tajbakhsh</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Shin</surname>
<given-names>JY</given-names>
</name>
<name>
<surname>Gurudu</surname>
<given-names>SR</given-names>
</name>
<name>
<surname>Todd Hurst</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Kendall</surname>
<given-names>CB</given-names>
</name>
<name>
<surname>Gotway</surname>
<given-names>MB</given-names>
</name>
<etal/>
</person-group> <article-title>Convolutional neural networks for medical image analysis: full training or fine tuning?</article-title> <source>IEEE Trans Med Imaging</source> (<year>2016</year>) <volume>35</volume>(<issue>5</issue>):<fpage>1299</fpage>&#x2013;<lpage>312</lpage>. <pub-id pub-id-type="doi">10.1109/tmi.2016.2535302</pub-id>
</citation>
</ref>
<ref id="B10">
<label>10.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jacob</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Wittek</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Pancotti</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Rebentrost</surname>
<given-names>P</given-names>
</name>
<name>
<surname>Wiebe</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Lloyd</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Quantum machine learning</article-title>. <source>Nature</source> (<year>2017</year>) <volume>549</volume>(<issue>7671</issue>):<fpage>195</fpage>&#x2013;<lpage>202</lpage>. <pub-id pub-id-type="doi">10.1038/nature23474</pub-id>
</citation>
</ref>
<ref id="B11">
<label>11.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schuld</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Killoran</surname>
<given-names>N</given-names>
</name>
</person-group>. <article-title>Quantum machine learning in feature hilbert spaces</article-title>. <source>Phys Rev Lett</source> (<year>2019</year>) <volume>122</volume>(<issue>4</issue>):<fpage>040504</fpage>. <pub-id pub-id-type="doi">10.1103/physrevlett.122.040504</pub-id>
</citation>
</ref>
<ref id="B12">
<label>12.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schuld</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Sinayskiy</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Petruccione</surname>
<given-names>F</given-names>
</name>
</person-group>. <article-title>An introduction to quantum machine learning</article-title>. <source>Contemp Phys</source> (<year>2015</year>) <volume>56</volume>(<issue>2</issue>):<fpage>172</fpage>&#x2013;<lpage>85</lpage>. <pub-id pub-id-type="doi">10.1080/00107514.2014.964942</pub-id>
</citation>
</ref>
<ref id="B13">
<label>13.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lloyd</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Mohseni</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Rebentrost</surname>
<given-names>P</given-names>
</name>
</person-group>. <article-title>Quantum algorithms for supervised and unsupervised machine learning</article-title>. <source>arXiv preprint arXiv:1307.0411</source> (<year>2013</year>). <pub-id pub-id-type="doi">10.48550/arXiv.1307.041</pub-id>
</citation>
</ref>
<ref id="B14">
<label>14.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Preskill</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Quantum computing in the nisq era and beyond</article-title>. <source>Quantum</source> (<year>2018</year>) <volume>2</volume>(<issue>79</issue>):<fpage>79</fpage>. <pub-id pub-id-type="doi">10.22331/q-2018-08-06-79</pub-id>
</citation>
</ref>
<ref id="B15">
<label>15.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bharti</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Cervera-Lierta</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Kyaw</surname>
<given-names>TH</given-names>
</name>
<name>
<surname>Haug</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Alperin-Lea</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Anand</surname>
<given-names>A</given-names>
</name>
<etal/>
</person-group> <article-title>Noisy intermediate-scale quantum algorithms</article-title>. <source>Rev Mod Phys</source> (<year>2022</year>) <volume>94</volume>(<issue>1</issue>):<fpage>015004</fpage>. <pub-id pub-id-type="doi">10.1103/revmodphys.94.015004</pub-id>
</citation>
</ref>
<ref id="B16">
<label>16.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Benedetti</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Lloyd</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Sack</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Fiorentini</surname>
<given-names>M</given-names>
</name>
</person-group>. <article-title>Parameterized quantum circuits as machine learning models</article-title>. <source>Quan Sci Technology</source> (<year>2019</year>) <volume>4</volume>(<issue>4</issue>):<fpage>043001</fpage>. <pub-id pub-id-type="doi">10.1088/2058-9565/ab4eb5</pub-id>
</citation>
</ref>
<ref id="B17">
<label>17.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sim</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Johnson</surname>
<given-names>PD</given-names>
</name>
<name>
<surname>Aspuru-Guzik</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Expressibility and entangling capability of parameterized quantum circuits for hybrid quantum-classical algorithms</article-title>. <source>Adv Quan Tech</source> (<year>2019</year>) <volume>2</volume>(<issue>12</issue>):<fpage>1900070</fpage>. <pub-id pub-id-type="doi">10.1002/qute.201900070</pub-id>
</citation>
</ref>
<ref id="B18">
<label>18.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cerezo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Arrasmith</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Babbush</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Benjamin</surname>
<given-names>SC</given-names>
</name>
<name>
<surname>Endo</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Fujii</surname>
<given-names>K</given-names>
</name>
<etal/>
</person-group> <article-title>Variational quantum algorithms</article-title>. <source>Nat Rev Phys</source> (<year>2021</year>) <volume>3</volume>(<issue>9</issue>):<fpage>625</fpage>&#x2013;<lpage>44</lpage>. <pub-id pub-id-type="doi">10.1038/s42254-021-00348-9</pub-id>
</citation>
</ref>
<ref id="B19">
<label>19.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>H-Y</given-names>
</name>
<name>
<surname>Kueng</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Preskill</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>Information-theoretic bounds on quantum advantage in machine learning</article-title>. <source>Phys Rev Lett</source> (<year>2021</year>) <volume>126</volume>(<issue>19</issue>):<fpage>190505</fpage>. <pub-id pub-id-type="doi">10.1103/physrevlett.126.190505</pub-id>
</citation>
</ref>
<ref id="B20">
<label>20.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Aharonov</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Cotler</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Xiao-Liang</surname>
<given-names>Q</given-names>
</name>
</person-group>. <article-title>Quantum algorithmic measurement</article-title>. <source>Nat Commun</source> (<year>2022</year>) <volume>13</volume>(<issue>1</issue>):<fpage>887</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-021-27922-0</pub-id>
</citation>
</ref>
<ref id="B21">
<label>21.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schuld</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Sweke</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Meyer</surname>
<given-names>JJ</given-names>
</name>
</person-group>. <article-title>Effect of data encoding on the expressive power of variational quantum-machine-learning models</article-title>. <source>Phys Rev A</source> (<year>2021</year>) <volume>103</volume>(<issue>3</issue>):<fpage>032430</fpage>. <pub-id pub-id-type="doi">10.1103/physreva.103.032430</pub-id>
</citation>
</ref>
<ref id="B22">
<label>22.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Caro</surname>
<given-names>MC</given-names>
</name>
<name>
<surname>Gil-Fuster</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Meyer</surname>
<given-names>JJ</given-names>
</name>
<name>
<surname>Eisert</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Sweke</surname>
<given-names>R</given-names>
</name>
</person-group>. <article-title>Encoding-dependent generalization bounds for parametrized quantum circuits</article-title>. <source>Quantum</source> (<year>2021</year>) <volume>5</volume>:<fpage>582</fpage>. <pub-id pub-id-type="doi">10.22331/q-2021-11-17-582</pub-id>
</citation>
</ref>
<ref id="B23">
<label>23.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Thanasilp</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Cerezo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Holmes</surname>
<given-names>Z</given-names>
</name>
</person-group>. <article-title>Exponential concentration and untrainability in quantum kernel methods</article-title>,. <comment>
<italic>arXiv preprint arXiv:2208.11060</italic>
</comment>(<year>2022</year>). <pub-id pub-id-type="doi">10.1038/s41467-024-49287-w</pub-id>
</citation>
</ref>
<ref id="B24">
<label>24.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Abbas</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Sutter</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Zoufal</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Lucchi</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Figalli</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Woerner</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>The power of quantum neural networks</article-title>. <source>Nat Comput Sci</source> (<year>2021</year>) <volume>1</volume>(<issue>6</issue>):<fpage>403</fpage>&#x2013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.1038/s43588-021-00084-1</pub-id>
</citation>
</ref>
<ref id="B25">
<label>25.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Caro</surname>
<given-names>MC</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>H-Y</given-names>
</name>
<name>
<surname>Cerezo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Sharma</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Sornborger</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Cincio</surname>
<given-names>L</given-names>
</name>
<etal/>
</person-group> <article-title>Generalization in quantum machine learning from few training data</article-title>. <source>Nat Commun</source> (<year>2022</year>) <volume>13</volume>(<issue>1</issue>):<fpage>4919</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-022-32550-3</pub-id>
</citation>
</ref>
<ref id="B26">
<label>26.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>McClean</surname>
<given-names>JR</given-names>
</name>
<name>
<surname>Boixo</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Smelyanskiy</surname>
<given-names>VN</given-names>
</name>
<name>
<surname>Babbush</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Neven</surname>
<given-names>H</given-names>
</name>
</person-group>. <article-title>Barren plateaus in quantum neural network training landscapes</article-title>. <source>Nat Commun</source> (<year>2018</year>) <volume>9</volume>(<issue>1</issue>):<fpage>4812</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-018-07090-4</pub-id>
</citation>
</ref>
<ref id="B27">
<label>27.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Larocca</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Thanasilp</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Sharma</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Jacob</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Coles</surname>
<given-names>PJ</given-names>
</name>
<etal/>
</person-group> <article-title>A review of barren plateaus in variational quantum computing</article-title>. <source>arXiv preprint arXiv:2405.00781</source> (<year>2024</year>). <pub-id pub-id-type="doi">10.48550/arXiv.2405.00781</pub-id>
</citation>
</ref>
<ref id="B28">
<label>28.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Holmes</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Sharma</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Cerezo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Coles</surname>
<given-names>PJ</given-names>
</name>
</person-group>. <article-title>Connecting ansatz expressibility to gradient magnitudes and barren plateaus</article-title>. <source>PRX Quan</source> (<year>2022</year>) <volume>3</volume>:<fpage>010313</fpage>. <pub-id pub-id-type="doi">10.1103/prxquantum.3.010313</pub-id>
</citation>
</ref>
<ref id="B29">
<label>29.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Grant</surname>
<given-names>E</given-names>
</name>
<name>
<surname>Benedetti</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Hallam</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Lockhart</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Stojevic</surname>
<given-names>V</given-names>
</name>
<etal/>
</person-group> <article-title>Hierarchical quantum classifiers</article-title>. <source>npj Quan Inf</source> (<year>2018</year>) <volume>4</volume>(<issue>1</issue>):<fpage>65</fpage>. <pub-id pub-id-type="doi">10.1038/s41534-018-0116-9</pub-id>
</citation>
</ref>
<ref id="B30">
<label>30.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pesah</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Cerezo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Tyler</surname>
<given-names>V</given-names>
</name>
<name>
<surname>Sornborger</surname>
<given-names>AT</given-names>
</name>
<name>
<surname>Coles</surname>
<given-names>PJ</given-names>
</name>
</person-group>. <article-title>Absence of barren plateaus in quantum convolutional neural networks</article-title>. <source>Phys Rev X</source> (<year>2021</year>) <volume>11</volume>(<issue>4</issue>):<fpage>041011</fpage>. <pub-id pub-id-type="doi">10.1103/physrevx.11.041011</pub-id>
</citation>
</ref>
<ref id="B31">
<label>31.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cong</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Choi</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Lukin</surname>
<given-names>MD</given-names>
</name>
</person-group>. <article-title>Quantum convolutional neural networks</article-title>. <source>Nat Phys</source> (<year>2019</year>) <volume>15</volume>(<issue>12</issue>):<fpage>1273</fpage>&#x2013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1038/s41567-019-0648-8</pub-id>
</citation>
</ref>
<ref id="B32">
<label>32.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hur</surname>
<given-names>T</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>DK</given-names>
</name>
</person-group>. <article-title>Quantum convolutional neural network for classical data classification</article-title>. <source>Quan Machine Intelligence</source> (<year>2022</year>) <volume>4</volume>(<issue>1</issue>):<fpage>3</fpage>. <pub-id pub-id-type="doi">10.1007/s42484-021-00061-x</pub-id>
</citation>
</ref>
<ref id="B33">
<label>33.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kim</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Huh</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>DK</given-names>
</name>
</person-group>. <article-title>Classical-to-quantum convolutional neural network transfer learning</article-title>. <source>Neurocomputing</source> (<year>2023</year>) <volume>555</volume>:<fpage>126643</fpage>. <pub-id pub-id-type="doi">10.1016/j.neucom.2023.126643</pub-id>
</citation>
</ref>
<ref id="B34">
<label>34.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lourens</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Sinayskiy</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>DK</given-names>
</name>
<name>
<surname>Blank</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Petruccione</surname>
<given-names>F</given-names>
</name>
</person-group>. <article-title>Hierarchical quantum circuit representations for neural architecture search</article-title>. <source>npj Quan Inf</source> (<year>2023</year>) <volume>9</volume>(<issue>1</issue>):<fpage>79</fpage>. <pub-id pub-id-type="doi">10.1038/s41534-023-00747-z</pub-id>
</citation>
</ref>
<ref id="B35">
<label>35.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Oh</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>DK</given-names>
</name>
</person-group>. <article-title>Quantum support vector data description for anomaly detection</article-title>. <source>Machine Learn Sci Technology</source> (<year>2024</year>) <volume>5</volume>:<fpage>035052</fpage>. <pub-id pub-id-type="doi">10.1088/2632-2153/ad6be8</pub-id>
</citation>
</ref>
<ref id="B36">
<label>36.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>Q</given-names>
</name>
<name>
<surname>Long</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Yuan</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>Y</given-names>
</name>
</person-group>. <article-title>Quantum convolutional neural network for image classification</article-title>. <source>Pattern Anal Appl</source> (<year>2023</year>) <volume>26</volume>(<issue>2</issue>):<fpage>655</fpage>&#x2013;<lpage>67</lpage>. <pub-id pub-id-type="doi">10.1007/s10044-022-01113-z</pub-id>
</citation>
</ref>
<ref id="B37">
<label>37.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Smaldone</surname>
<given-names>AM</given-names>
</name>
<name>
<surname>Kyro</surname>
<given-names>GW</given-names>
</name>
<name>
<surname>Batista</surname>
<given-names>VS</given-names>
</name>
</person-group>. <article-title>Quantum convolutional neural networks for multi-channel supervised learning</article-title>. <source>Quan Machine Intelligence</source> (<year>2023</year>) <volume>5</volume>(<issue>2</issue>):<fpage>41</fpage>. <pub-id pub-id-type="doi">10.1007/s42484-023-00130-3</pub-id>
</citation>
</ref>
<ref id="B38">
<label>38.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Banchi</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Pereira</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Pirandola</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Generalization in quantum machine learning: a quantum information standpoint</article-title>. <source>PRX Quan</source> (<year>2021</year>) <volume>2</volume>:<fpage>040321</fpage>. <pub-id pub-id-type="doi">10.1103/prxquantum.2.040321</pub-id>
</citation>
</ref>
<ref id="B39">
<label>39.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Goodfellow</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Bengio</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Courville</surname>
<given-names>A</given-names>
</name>
</person-group>. <source>Deep learning</source>. <publisher-name>MIT Press</publisher-name> (<year>2016</year>). <comment>Available from: <ext-link ext-link-type="uri" xlink:href="http://www.deeplearningbook.org">http://www.deeplearningbook.org</ext-link>.</comment>
</citation>
</ref>
<ref id="B40">
<label>40.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Bengio</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Hardt</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Benjamin</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Vinyals</surname>
<given-names>O</given-names>
</name>
</person-group>. <article-title>Understanding deep learning (still) requires rethinking generalization</article-title>. <source>Commun ACM</source> (<year>2021</year>) <volume>64</volume>(<issue>3</issue>):<fpage>107</fpage>&#x2013;<lpage>15</lpage>. <pub-id pub-id-type="doi">10.1145/3446776</pub-id>
</citation>
</ref>
<ref id="B41">
<label>41.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lu</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Lu</surname>
<given-names>J</given-names>
</name>
</person-group>. <article-title>A universal approximation theorem of deep neural networks for expressing probability distributions</article-title>. <source>Adv Neural Inf Process Syst</source> (<year>2020</year>) <volume>33</volume>:<fpage>3094</fpage>&#x2013;<lpage>105</lpage>.</citation>
</ref>
<ref id="B42">
<label>42.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ruder</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>An overview of gradient descent optimization algorithms</article-title>. <source>arXiv preprint arXiv:1609.04747</source> (<year>2016</year>). <pub-id pub-id-type="doi">10.48550/arXiv.1609.04747</pub-id>
</citation>
</ref>
<ref id="B43">
<label>43.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mitarai</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Negoro</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Kitagawa</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Fujii</surname>
<given-names>K</given-names>
</name>
</person-group>. <article-title>Quantum circuit learning</article-title>. <source>Phys Rev A</source> (<year>2018</year>) <volume>98</volume>(<issue>3</issue>):<fpage>032309</fpage>. <pub-id pub-id-type="doi">10.1103/physreva.98.032309</pub-id>
</citation>
</ref>
<ref id="B44">
<label>44.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schuld</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Bergholm</surname>
<given-names>V</given-names>
</name>
<name>
<surname>Gogolin</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Izaac</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Killoran</surname>
<given-names>N</given-names>
</name>
</person-group>. <article-title>Evaluating analytic gradients on quantum hardware</article-title>. <source>Phys Rev A</source> (<year>2019</year>) <volume>99</volume>(<issue>3</issue>):<fpage>032331</fpage>. <pub-id pub-id-type="doi">10.1103/physreva.99.032331</pub-id>
</citation>
</ref>
<ref id="B45">
<label>45.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>MacCormack</surname>
<given-names>I</given-names>
</name>
<name>
<surname>Delaney</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Galda</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Aggarwal</surname>
<given-names>N</given-names>
</name>
<name>
<surname>Narang</surname>
<given-names>P</given-names>
</name>
</person-group>. <article-title>Branching quantum convolutional neural networks</article-title>. <source>Phys Rev Res</source> (<year>2022</year>) <volume>4</volume>(<issue>1</issue>):<fpage>013117</fpage>. <pub-id pub-id-type="doi">10.1103/physrevresearch.4.013117</pub-id>
</citation>
</ref>
<ref id="B46">
<label>46.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mari</surname>
<given-names>A</given-names>
</name>
<name>
<surname>Bromley</surname>
<given-names>TR</given-names>
</name>
<name>
<surname>Izaac</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Schuld</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Killoran</surname>
<given-names>N</given-names>
</name>
</person-group>. <article-title>Transfer learning in hybrid classical-quantum neural networks</article-title>. <source>Quantum</source> (<year>2020</year>) <volume>4</volume>:<fpage>340</fpage>. <pub-id pub-id-type="doi">10.22331/q-2020-10-09-340</pub-id>
</citation>
</ref>
<ref id="B47">
<label>47.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z</given-names>
</name>
<name>
<surname>Han</surname>
<given-names>R</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Shang</surname>
<given-names>R</given-names>
</name>
<etal/>
</person-group> <article-title>Quantum recurrent neural networks for sequential learning</article-title>. <source>Neural Networks</source> (<year>2023</year>) <volume>166</volume>:<fpage>148</fpage>&#x2013;<lpage>61</lpage>. <pub-id pub-id-type="doi">10.1016/j.neunet.2023.07.003</pub-id>
</citation>
</ref>
<ref id="B48">
<label>48.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lloyd</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Weedbrook</surname>
<given-names>C</given-names>
</name>
</person-group>. <article-title>Quantum generative adversarial learning</article-title>. <source>Phys Rev Lett</source> (<year>2018</year>) <volume>121</volume>:<fpage>040502</fpage>. <pub-id pub-id-type="doi">10.1103/physrevlett.121.040502</pub-id>
</citation>
</ref>
<ref id="B49">
<label>49.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dallaire-Demers</surname>
<given-names>P-L</given-names>
</name>
<name>
<surname>Killoran</surname>
<given-names>N</given-names>
</name>
</person-group>. <article-title>Quantum generative adversarial networks</article-title>. <source>Phys Rev A</source> (<year>2018</year>) <volume>98</volume>:<fpage>012324</fpage>. <pub-id pub-id-type="doi">10.1103/physreva.98.012324</pub-id>
</citation>
</ref>
<ref id="B50">
<label>50.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yen-Chi Chen</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Wei</surname>
<given-names>T-C</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Yoo</surname>
<given-names>S</given-names>
</name>
</person-group>. <article-title>Hybrid quantum-classical graph convolutional network</article-title>. <source>arXiv preprint arXiv:2101.06189</source> (<year>2021</year>). <pub-id pub-id-type="doi">10.48550/arXiv.2101.06189</pub-id>
</citation>
</ref>
<ref id="B51">
<label>51.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>T&#xfc;ys&#xfc;z</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Rieger</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Novotny</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Demirk&#xf6;z</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Dobos</surname>
<given-names>D</given-names>
</name>
<name>
<surname>Potamianos</surname>
<given-names>K</given-names>
</name>
<etal/>
</person-group> <article-title>Hybrid quantum classical graph neural networks for particle track reconstruction</article-title>. <source>Quan Machine Intelligence</source> (<year>2021</year>) <volume>3</volume>(<issue>2</issue>):<fpage>29</fpage>. <pub-id pub-id-type="doi">10.1007/s42484-021-00055-9</pub-id>
</citation>
</ref>
<ref id="B52">
<label>52.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Park</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Huh</surname>
<given-names>J</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>DK</given-names>
</name>
</person-group>. <article-title>Variational quantum one-class classifier</article-title>. <source>Machine Learn Sci Technology</source> (<year>2023</year>) <volume>4</volume>(<issue>1</issue>):<fpage>015006</fpage>. <pub-id pub-id-type="doi">10.1088/2632-2153/acafd5</pub-id>
</citation>
</ref>
<ref id="B53">
<label>53.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jaderberg</surname>
<given-names>B</given-names>
</name>
<name>
<surname>Anderson</surname>
<given-names>LW</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>W</given-names>
</name>
<name>
<surname>Albanie</surname>
<given-names>S</given-names>
</name>
<name>
<surname>Kiffner</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Jaksch</surname>
<given-names>D</given-names>
</name>
</person-group>. <article-title>Quantum self-supervised learning</article-title>. <source>Quan Sci Technology</source> (<year>2022</year>) <volume>7</volume>(<issue>3</issue>):<fpage>035005</fpage>. <pub-id pub-id-type="doi">10.1088/2058-9565/ac6825</pub-id>
</citation>
</ref>
<ref id="B54">
<label>54.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>LeCun</surname>
<given-names>Y</given-names>
</name>
<name>
<surname>Cortes</surname>
<given-names>C</given-names>
</name>
<name>
<surname>Burges</surname>
<given-names>CJ</given-names>
</name>
</person-group>. <article-title>Mnist handwritten digit database</article-title>. <source>ATT Labs</source> (<year>2010</year>). <comment>Available from: <ext-link ext-link-type="uri" xlink:href="http://yann.lecun.com/exdb/mnist">http://yann.lecun.com/exdb/mnist</ext-link>,</comment>
<volume>2</volume>.</citation>
</ref>
<ref id="B55">
<label>55.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Srinivasan</surname>
<given-names>A</given-names>
</name>
</person-group>. <article-title>Statlog (Landsat satellite)</article-title>. <source>UCI Machine Learn Repository</source> (<year>1993</year>). <pub-id pub-id-type="doi">10.24432/C55887</pub-id>
</citation>
</ref>
<ref id="B56">
<label>56.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xiao</surname>
<given-names>H</given-names>
</name>
<name>
<surname>Rasul</surname>
<given-names>K</given-names>
</name>
<name>
<surname>Vollgraf</surname>
<given-names>R</given-names>
</name>
</person-group>. <article-title>Fashion-mnist: a novel image dataset for benchmarking machine learning algorithms</article-title>. <source>CoRR</source> (<year>2017</year>) <fpage>07747</fpage>. <comment>abs/1708</comment>. <pub-id pub-id-type="doi">10.48550/arXiv.1708.07747</pub-id>
</citation>
</ref>
<ref id="B57">
<label>57.</label>
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Sigillito</surname>
<given-names>V</given-names>
</name>
<name>
<surname>Baker</surname>
<given-names>K</given-names>
</name>
</person-group>. <article-title>Ionosphere</article-title>. In: <source>UCI machine learning repository</source> (<year>1989</year>). <pub-id pub-id-type="doi">10.24432/C5W01B</pub-id>
</citation>
</ref>
<ref id="B58">
<label>58.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vatan</surname>
<given-names>F</given-names>
</name>
<name>
<surname>Williams</surname>
<given-names>C</given-names>
</name>
</person-group>. <article-title>Optimal quantum circuits for general two-qubit gates</article-title>. <source>Phys Rev A</source> (<year>2004</year>) <volume>69</volume>(<issue>3</issue>):<fpage>032315</fpage>. <pub-id pub-id-type="doi">10.1103/physreva.69.032315</pub-id>
</citation>
</ref>
<ref id="B59">
<label>59.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Diederik</surname>
<given-names>PK</given-names>
</name>
<name>
<surname>Jimmy</surname>
<given-names>B</given-names>
</name>
</person-group>. <article-title>Adam: a method for stochastic optimization</article-title>. <source>arXiv preprint arXiv:1412.6980</source> (<year>2014</year>). <pub-id pub-id-type="doi">10.48550/arXiv.1412.6980</pub-id>
</citation>
</ref>
<ref id="B60">
<label>60.</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cerezo</surname>
<given-names>M</given-names>
</name>
<name>
<surname>Verdon</surname>
<given-names>G</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>H-Y</given-names>
</name>
<name>
<surname>Cincio</surname>
<given-names>L</given-names>
</name>
<name>
<surname>Coles</surname>
<given-names>PJ</given-names>
</name>
</person-group>. <article-title>Challenges and opportunities in quantum machine learning</article-title>. <source>Nat Comput Sci</source> (<year>2022</year>) <volume>2</volume>:<fpage>567</fpage>&#x2013;<lpage>76</lpage>. <pub-id pub-id-type="doi">10.1038/s43588-022-00311-3</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>