<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Quantum Sci. Technol.</journal-id>
<journal-title>Frontiers in Quantum Science and Technology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Quantum Sci. Technol.</abbrev-journal-title>
<issn pub-type="epub">2813-2181</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1200975</article-id>
<article-id pub-id-type="doi">10.3389/frqst.2023.1200975</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Quantum Science and Technology</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Similarity-based parameter transferability in the quantum approximate optimization algorithm</article-title>
<alt-title alt-title-type="left-running-head">Galda et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/frqst.2023.1200975">10.3389/frqst.2023.1200975</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Galda</surname>
<given-names>Alexey</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<xref ref-type="fn" rid="fn1">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1886674/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Gupta</surname>
<given-names>Eesh</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="fn" rid="fn1">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2319723/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Falla</surname>
<given-names>Jose</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2209597/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Liu</surname>
<given-names>Xiaoyuan</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2352946/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Lykov</surname>
<given-names>Danylo</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Alexeev</surname>
<given-names>Yuri</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2189350/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Safro</surname>
<given-names>Ilya</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2091200/overview"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>James Franck Institute</institution>, <institution>University of Chicago</institution>, <addr-line>Chicago</addr-line>, <addr-line>IL</addr-line>, <country>United States</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Computational Science Division</institution>, <institution>Argonne National Laboratory</institution>, <addr-line>Lemont</addr-line>, <addr-line>IL</addr-line>, <country>United States</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Department of Physics and Astronomy</institution>, <institution>Rutgers University</institution>, <addr-line>Piscataway</addr-line>, <addr-line>NJ</addr-line>, <country>United States</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Department of Physics and Astronomy</institution>, <institution>University of Delaware</institution>, <addr-line>Newark</addr-line>, <addr-line>DE</addr-line>, <country>United States</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Department of Computer and Information Sciences</institution>, <institution>University of Delaware</institution>, <addr-line>Newark</addr-line>, <addr-line>DE</addr-line>, <country>United States</country>
</aff>
<aff id="aff6">
<sup>6</sup>
<institution>Fujitsu Research of America, Inc.</institution>, <addr-line>Sunnyvale</addr-line>, <addr-line>CA</addr-line>, <country>United States</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1131825/overview">Laszlo Gyongyosi</ext-link>, Budapest University of Technology and Economics, Hungary</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1556707/overview">Mingxing Luo</ext-link>, Southwest Jiaotong University, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/81712/overview">Prasanta Panigrahi</ext-link>, Indian Institute of Science Education and Research Kolkata, India</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Alexey Galda, <email>alex.galda@gmail.com</email>
</corresp>
<fn fn-type="equal" id="fn1">
<label>
<sup>&#x2020;</sup>
</label>
<p>These authors have contributed equally to this work and share first authorship</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>13</day>
<month>07</month>
<year>2023</year>
</pub-date>
<pub-date pub-type="collection">
<year>2023</year>
</pub-date>
<volume>2</volume>
<elocation-id>1200975</elocation-id>
<history>
<date date-type="received">
<day>05</day>
<month>04</month>
<year>2023</year>
</date>
<date date-type="accepted">
<day>15</day>
<month>06</month>
<year>2023</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2023 Galda, Gupta, Falla, Liu, Lykov, Alexeev and Safro.</copyright-statement>
<copyright-year>2023</copyright-year>
<copyright-holder>Galda, Gupta, Falla, Liu, Lykov, Alexeev and Safro</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>The quantum approximate optimization algorithm (QAOA) is one of the most promising candidates for achieving quantum advantage through quantum-enhanced combinatorial optimization. A near-optimal solution to the combinatorial optimization problem is achieved by preparing a quantum state through the optimization of quantum circuit parameters. Optimal QAOA parameter concentration effects for special MaxCut problem instances have been observed, but a rigorous study of the subject is still lacking. In this work we show clustering of optimal QAOA parameters around specific values; consequently, successful transferability of parameters between different QAOA instances can be explained and predicted based on local properties of the graphs, including the type of subgraphs (lightcones) from which graphs are composed as well as the overall degree of nodes in the graph (parity). We apply this approach to several instances of random graphs with a varying number of nodes as well as parity and show that one can use optimal donor graph QAOA parameters as near-optimal parameters for larger acceptor graphs with comparable approximation ratios. This work presents a pathway to identifying classes of combinatorial optimization instances for which variational quantum algorithms such as QAOA can be substantially accelerated.</p>
</abstract>
<kwd-group>
<kwd>quantum computing</kwd>
<kwd>quantum software</kwd>
<kwd>quantum optimization</kwd>
<kwd>quantum approximate optimization algorithm</kwd>
<kwd>parameter transferability</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Quantum Information Theory</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>Quantum computing seeks to exploit the quantum mechanical concepts of entanglement and superposition to perform a computation that is significantly faster and more efficient than what can be achieved by using the most powerful supercomputers available today (<xref ref-type="bibr" rid="B22">Preskill, 2018</xref>; <xref ref-type="bibr" rid="B4">Arute et al., 2019</xref>). Demonstrating quantum advantage with optimization algorithms (<xref ref-type="bibr" rid="B2">Alexeev et al., 2021</xref>) is poised to have a broad impact on science and humanity by allowing us to solve problems on a global scale, including finance (<xref ref-type="bibr" rid="B12">Herman et al., 2022</xref>), biology (<xref ref-type="bibr" rid="B21">Outeiral et al., 2021</xref>), and energy (<xref ref-type="bibr" rid="B15">Joseph et al., 2023</xref>). Variational quantum algorithms, a class of hybrid quantum-classical algorithms, are considered primary candidates for such tasks and consist of parameterized quantum circuits with parameters updated in classical computation. The quantum approximate optimization algorithm (QAOA) (<xref ref-type="bibr" rid="B14">Hogg, 2000</xref>; <xref ref-type="bibr" rid="B13">Hogg and Portnov, 2000</xref>; <xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>; <xref ref-type="bibr" rid="B11">Hadfield et al., 2019</xref>) is a variational algorithm for solving classical combinatorial optimization problems. In the domain of optimization on graphs, it is most often used to solve NP-hard problems such as MaxCut (<xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>), community detection (<xref ref-type="bibr" rid="B26">Shaydulin et al., 2019c</xref>), and partitioning (<xref ref-type="bibr" rid="B28">Ushijima-Mwesigwa et al., 2021</xref>) by mapping them onto a classical spin-glass model (also known as the Ising model) and minimizing the corresponding energy, a task that in itself is NP-hard.</p>
<p>In this work we demonstrate two related key elements of optimal QAOA parameter transferability. First, by analyzing the distributions of subgraphs from two QAOA MaxCut instance graphs, one can predict how close the optimized QAOA parameters for one instance are to the optimal QAOA parameters for another. Second, by analyzing the overall parity of both donor-acceptor pairs, one can predict good transferability between those QAOA MaxCut instances. The measure of transferability of optimized parameters between MaxCut QAOA instances on two graphs can be expressed through the value of the approximation ratio, which is defined as the ratio of the energy of the corresponding QAOA circuit, evaluated with the optimized parameters <italic>&#x3b3;</italic>, <italic>&#x3b2;</italic>, divided by the energy of the optimal MaxCut solution for the graph.</p>
<p>While the optimal solution is not known in general for relatively small instances (graphs with up to 256 nodes are considered in this paper), it can be found by using classical algorithms, such as the Gurobi solver (<xref ref-type="bibr" rid="B10">Gurobi Optimization, 2021</xref>)<xref ref-type="fn" rid="fn2">
<sup>1</sup>
</xref>. We first focus our attention on similarity based on the subgraph decomposition of random graphs and show that good transferability of optimized parameters between two graphs is directly determined by the transferability between all possible permutations of pairs of individual subgraphs. The relevant subgraphs of these graphs are defined by the QAOA quantum circuit depth parameter <italic>p</italic>. In this work we focus on the case <italic>p</italic> &#x3d; 1; however, our approach can be extended to larger values of <italic>p</italic>. Higher values of <italic>p</italic> lead to an increasing number of subgraphs to be considered, but the general idea of the approach remains the same. This question is beyond the scope of this paper and will be addressed in our future work. We then move to similarity based on graph parity and determine that we can predict good optimal parameter transferability between donor-acceptor graph pairs with similar parities. Here, too, more work remains to be done regarding the structural effects of graphs on optimal parameter transferability.</p>
<p>Based on the analysis of the mutual transferability of optimized QAOA parameters between all relevant subgraphs for computing the MaxCut cost function of random graphs, we show good transferability <italic>within</italic> the classes of odd and even random graphs of arbitrary size. We also show that transferability is poor <italic>between</italic> the classes of even and odd random graphs, in both directions, based on the poor transferability of the optimized QAOA parameters between the subgraphs of the corresponding graphs. When considering the most general case of arbitrary random graphs, we construct the transferability map between all possible subgraphs of such graphs, with an upper limit of node connectivity <italic>d</italic>
<sub>max</sub> &#x3d; 6, and use it to demonstrate that in order to find optimized parameters for a MaxCut QAOA instance on a large 64-, 128-, or 256-node random graph, under specific conditions, one can reuse the optimized parameters from a random graph of a much smaller size, <italic>N</italic> &#x3d; 6, with only a &#x223c;1% reduction in the approximation ratio.</p>
<p>This paper is structured as follows. In <xref ref-type="sec" rid="s2">Section 2</xref> we present the relevant background material on QAOA. In <xref ref-type="sec" rid="s3">Section 3</xref> we consider optimized QAOA parameter transferability properties between all possible subgraphs of random graphs of degree up to <italic>d</italic>
<sub>max</sub> &#x3d; 6. We then extend the consideration to parameter transferability using graph parity as a metric, and we demonstrate the power of the proposed approach by performing optimal transferability of QAOA parameters in many instances of donor-acceptor graph pairs of differing sizes and parity. We find that one can effectively transfer optimal parameters from smaller donor graphs to larger acceptor graphs, using similarities based on subgraph decomposition and parity as indicators of good transferability. In <xref ref-type="sec" rid="s4">Section 4</xref> we conclude with a summary of our results and an outlook on future advances with our approach.</p>
</sec>
<sec id="s2">
<title>2 QAOA</title>
<p>The quantum approximate optimization algorithm is a hybrid quantum-classical algorithm that combines a parameterized quantum evolution with a classical outer-loop optimizer to approximately solve binary optimization problems (<xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>; <xref ref-type="bibr" rid="B11">Hadfield et al., 2019</xref>). QAOA consists of <italic>p</italic> layers (also known as the circuit depth) of pairs of alternating operators, with each additional layer increasing the quality of the solution, assuming perfect noiseless execution of the corresponding quantum circuit. With quantum error correction not currently supported by modern quantum processors, practical implementations of QAOA are limited to <italic>p</italic> &#x2264; 3 because of noise and limited coherence of quantum devices imposing strict limitations on the circuit depth (<xref ref-type="bibr" rid="B31">Zhou et al., 2020</xref>). Motivated by the practical relevance of results, we focus on the case <italic>p</italic> &#x3d; 1 in this paper.</p>
<sec id="s2-1">
<title>2.1 QAOA background</title>
<p>Consider a combinatorial problem defined on a space of binary strings of length <italic>N</italic> that has <italic>m</italic> clauses. Each clause is a constraint satisfied by some assignment of the bit string. The objective function can be written as <inline-formula id="inf1">
<mml:math id="m1">
<mml:mi>C</mml:mi>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mo movablelimits="false" form="prefix">&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>&#x3b1;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>m</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:msub>
<mml:mrow>
<mml:mi>C</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>&#x3b1;</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula>, where <italic>z</italic> &#x3d; <italic>z</italic>
<sub>1</sub>
<italic>z</italic>
<sub>2</sub>&#x22ef;<italic>z</italic>
<sub>
<italic>N</italic>
</sub> is the bit string and <italic>C</italic>
<sub>
<italic>&#x3b1;</italic>
</sub>(<italic>z</italic>) &#x3d; 1 if <italic>z</italic> satisfies the clause <italic>&#x3b1;</italic>, and 0 otherwise. QAOA maps the combinatorial optimization problem onto a 2<sup>
<italic>N</italic>
</sup>-dimensional Hilbert space with computational basis vectors &#x7c;<italic>z</italic>&#x27e9; and encodes <italic>C</italic>(<italic>z</italic>) as an operator <italic>C</italic> diagonal in the computational basis.</p>
<p>At each call to the quantum computer, a trial state is prepared by applying a sequence of alternating quantum operators<disp-formula id="e1">
<mml:math id="m2">
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo>,</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo stretchy="false">&#x232a;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2254;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>B</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x2026;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>B</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:msub>
<mml:mrow>
<mml:mi>U</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo stretchy="false">&#x232a;</mml:mo>
<mml:mo>,</mml:mo>
</mml:math>
<label>(1)</label>
</disp-formula>where <italic>U</italic>
<sub>
<italic>C</italic>
</sub>(<italic>&#x3b3;</italic>) &#x3d; <italic>e</italic>
<sup>&#x2212;<italic>i&#x3b3;C</italic>
</sup> is the phase operator; <italic>U</italic>
<sub>
<italic>B</italic>
</sub>(<italic>&#x3b2;</italic>) &#x3d; <italic>e</italic>
<sup>&#x2212;<italic>i&#x3b2;B</italic>
</sup> is the mixing operator, with <italic>B</italic> defined as the operator of all single-bit <italic>&#x3c3;</italic>
<sup>
<italic>x</italic>
</sup> operators; <inline-formula id="inf2">
<mml:math id="m3">
<mml:mi>B</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mo movablelimits="false" form="prefix">&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:msubsup>
<mml:mrow>
<mml:mi>&#x3c3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msubsup>
</mml:math>
</inline-formula>; and &#x7c;<italic>s</italic>&#x27e9; is some easy-to-prepare initial state, usually taken to be the uniform superposition product state. The parameterized quantum circuit (1) is called the QAOA <italic>ansatz</italic>. We refer to the number of alternating operator pairs <italic>p</italic> as the QAOA <italic>depth</italic>. The selected parameters <inline-formula id="inf3">
<mml:math id="m4">
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo>,</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula> are said to define a <italic>schedule</italic>, analogous to a similar choice in quantum annealing.</p>
<p>Preparation of the state (1) is followed by a measurement in the computational basis. The output of repeated state preparation and measurement may be used by a classical outer-loop algorithm to select the schedule <inline-formula id="inf4">
<mml:math id="m5">
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo>,</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula>. We consider optimizing the expectation value of the objective function<disp-formula id="equ1">
<mml:math id="m6">
<mml:msub>
<mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">&#x27e8;</mml:mo>
<mml:mrow>
<mml:mi>C</mml:mi>
</mml:mrow>
<mml:mo stretchy="false">&#x27e9;</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mo stretchy="false">&#x2329;</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo>,</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mo stretchy="false">&#x7c;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mi>C</mml:mi>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mo>&#x20d7;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo stretchy="false">&#x232a;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
</mml:math>
</disp-formula>as originally proposed in (<xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>). The output of the overall procedure is the best bit string <italic>z</italic> found for the given combinatorial optimization problem. <xref ref-type="fig" rid="F1">Figure 1</xref> presents a schematic pipeline of the QAOA algorithm. We emphasize that the task of finding good QAOA parameters is challenging in general, for example, because of encountering barren plateaus (<xref ref-type="bibr" rid="B29">Wang et al., 2021</xref>; <xref ref-type="bibr" rid="B3">Anschuetz and Kiani, 2022</xref>). Acceleration of the optimal parameters search for a given QAOA depth <italic>p</italic> is the focus of many approaches aimed at demonstrating quantum advantage. Examples include warm- and multistart optimization (<xref ref-type="bibr" rid="B24">Shaydulin et al., 2019a</xref>; <xref ref-type="bibr" rid="B8">Egger et al., 2020</xref>), problem decomposition (<xref ref-type="bibr" rid="B25">Shaydulin et al., 2019b</xref>), instance structure analysis (<xref ref-type="bibr" rid="B23">Shaydulin et al., 2021</xref>), and parameter learning (<xref ref-type="bibr" rid="B17">Khairy et al., 2020</xref>).</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Schematic pipeline of a QAOA circuit. A parametrized ansatz is initialized, followed by series of applied unitaries that define the depth of the circuit. Finally, measurements are made in the computational basis, and the variational angles are classically optimized. This hybrid quantum-classical loop continues until convergence to an approximate solution is achieved.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g001.tif"/>
</fig>
</sec>
<sec id="s2-2">
<title>2.2 MaxCut</title>
<p>For studying the transferability of optimized QAOA parameters, we consider the MaxCut combinatorial optimization problem. Given an unweighted undirected simple graph <italic>G</italic> &#x3d; (<italic>V</italic>, <italic>E</italic>), the goal of the MaxCut problem is to find a partition of the graph&#x2019;s vertices into two complementary sets such that the number of edges between the two sets is maximized. In order to encode the problem in the QAOA setting, the input is a graph with &#x7c;<italic>V</italic>&#x7c; &#x3d; <italic>N</italic> vertices and &#x7c;<italic>E</italic>&#x7c; &#x3d; <italic>m</italic> edges, and the goal is to find a bit string <italic>z</italic> that maximizes<disp-formula id="e2">
<mml:math id="m7">
<mml:mi>C</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mi>k</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mi>E</mml:mi>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mrow>
<mml:mi>C</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
</mml:math>
<label>(2)</label>
</disp-formula>where<disp-formula id="equ2">
<mml:math id="m8">
<mml:msub>
<mml:mrow>
<mml:mi>C</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mi>&#x3c3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:msubsup>
<mml:mrow>
<mml:mi>&#x3c3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>k</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>z</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
<mml:mo>.</mml:mo>
</mml:math>
</disp-formula>
</p>
<p>It has been shown in (<xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>) that on a 3-regular graph, QAOA with <italic>p</italic> &#x3d; 1 produces a solution with an approximation ratio of at least 0.6924.</p>
</sec>
<sec id="s2-3">
<title>2.3 QAOA simulator and classical MaxCut solver</title>
<p>Calculating the approximation ratio for a particular MaxCut problem instance requires the optimal solution of the combinatorial optimization problem. This problem is known to be NP-hard, and classical solvers require exponential time to converge. For our experiments, we use the Gurobi solver (<xref ref-type="bibr" rid="B10">Gurobi Optimization, 2021</xref>) with the default configuration parameters, running the solver until it converges to the optimal solution. For our QAOA simulations, we use QTensor (<xref ref-type="bibr" rid="B18">Lykov et al., 2021</xref>), a large-scale quantum circuit simulator with step-dependent parallelization. QTensor simulates circuits based on a tensor network approach, and as such, it can provide an efficient approximation to certain classes of quantum states (<xref ref-type="bibr" rid="B6">Biamonte and Bergholm, 2017</xref>; <xref ref-type="bibr" rid="B16">Kardashin et al., 2021</xref>).</p>
</sec>
</sec>
<sec id="s3">
<title>3 Parameter transferability</title>
<p>Solving a QAOA instance calls for two types of executions of quantum circuits: iterative optimization of the QAOA parameters and the final sampling from the output state prepared with those parameters. While the latter is known to be impossible to simulate efficiently for large enough instances using classical hardware instead of a quantum processor (<xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>), the iterative energy calculation for the QAOA circuit during the classical optimization loop can be efficiently performed by using tensor network simulators for instances of a wide range of sizes (<xref ref-type="bibr" rid="B19">Lykov et al., 2020</xref>), as described in the preceding section. This is achieved by implementing considerable simplifications in how the expectation value of the problem Hamiltonian is calculated by employing a mathematical reformulation based on the notion of the reverse causal cone introduced in the seminal QAOA paper (<xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>). Moreover, in some instances, the entire search of the optimal parameters for a particular QAOA instance can be circumvented by reusing the optimized parameters from a different &#x201c;related&#x201d; instance, for example, for which the optimal parameters are concentrated in the same region.</p>
<p>Optimizing QAOA parameters for a relatively small graph, called the donor, and using them to prepare the QAOA state that maximizes the expectation value &#x27e8;<italic>C</italic>&#x27e9;<sub>
<italic>p</italic>
</sub> for the same problem on a larger graph, called the acceptor, is what we define as <italic>successful optimal parameter transferability</italic>, or just transferability of parameters, for brevity. The transferred parameters can be used either directly without change, as implemented in this paper, or as a &#x201c;warm start&#x201d; for further optimization. In either case, the high computational cost of optimizing the QAOA parameters, which grows rapidly as the QAOA depth <italic>p</italic> and the problem size are increased, can be significantly reduced. This approach presents a new direction for dramatically reducing the overall runtime of QAOA.</p>
<p>Optimal QAOA parameter concentration effects have been reported for several special cases, mainly focusing on random 3-regular graphs (<xref ref-type="bibr" rid="B7">Brandao et al., 2018</xref>; <xref ref-type="bibr" rid="B27">Streif and Leib, 2020</xref>; <xref ref-type="bibr" rid="B1">Akshay et al., 2021</xref>). <xref ref-type="bibr" rid="B7">Brandao et al. (2018)</xref> observed that the optimized QAOA parameters for the MaxCut problem obtained for a 3-regular graph are also nearly optimal for all other 3-regular graphs. In particular, the authors noted that in the limit of large <italic>N</italic>, where <italic>N</italic> is the number of nodes, the fraction of tree graphs asymptotically approaches 1. We note that, for example, in the sparse Erd&#xf6;s&#x2013;R&#xe9;nyi graphs, the trees are observed in short-distance neighborhoods with very high probability (<xref ref-type="bibr" rid="B20">Newman, 2018</xref>). As a result, in this limit, the objective function is the same for all 3-regular graphs, up to order 1/<italic>N</italic>.</p>
<p>The central question of this manuscript is under what conditions the optimized QAOA parameters for one graph also maximize the QAOA objective function for another graph. To answer that question, we study transferability between subgraphs of a graph, since the QAOA objective function is fully determined by the corresponding subgraphs of the instance graph, as well as transferability between graphs of similar parities, in order to determine structural effects of graphs on effective transferability.</p>
<sec id="s3-1">
<title>3.1 Subgraph transferability analysis</title>
<p>It was shown in the seminal QAOA paper (<xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>) that the expectation value of the QAOA objective function, &#x27e8;<italic>C</italic>&#x27e9;<sub>
<italic>p</italic>
</sub>, can be evaluated as a sum over contributions from subgraphs of the original graph, provided its degree is bounded and the diameter is larger than 2<italic>p</italic> (otherwise, the subgraphs cover the entire graph itself). The contributing subgraphs can be constructed by iterating over all edges of the original graph and selecting only the nodes that are <italic>p</italic> edges away from the edge. Through this process, any graph can be deconstructed into a set of subgraphs for a given <italic>p</italic>, and only those subgraphs contribute to &#x27e8;<italic>C</italic>&#x27e9;<sub>
<italic>p</italic>
</sub>, as also discussed in <xref ref-type="sec" rid="s2">Section 2</xref>.</p>
<p>We begin by analyzing the case of MaxCut instances on 3-regular random graphs for QAOA circuit depth <italic>p</italic> &#x3d; 1, which have three possible subgraphs (<xref ref-type="bibr" rid="B9">Farhi et al., 2014</xref>; <xref ref-type="bibr" rid="B7">Brandao et al., 2018</xref>). <xref ref-type="fig" rid="F2">Figure 2</xref> (top row) shows the landscapes of energy contributions from these subgraphs, evaluated for a range of <italic>&#x3b3;</italic> and <italic>&#x3b2;</italic> parameters. We can see that all maxima are located in the approximate vicinity of each other. As a result, the parameters optimized for any of the three graphs will also be near-optimal for the other two. Because any random 3-regular graph can be decomposed into these three subgraphs, for QAOA with <italic>p</italic> &#x3d; 1, this guarantees that optimized QAOA parameters can be successfully transferred between any two 3-regular random graphs, which is in full agreement with (<xref ref-type="bibr" rid="B7">Brandao et al., 2018</xref>).</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Landscapes of energy contributions for individual subgraphs of 3- (top row), 4- (middle row), and 5-regular (bottom row) random graphs, as a function of QAOA parameters <italic>&#x3b2;</italic> &#x2208;(0, <italic>&#x3c0;</italic>) and <italic>&#x3b3;</italic> &#x2208; (0, 2<italic>&#x3c0;</italic>). All subgraphs of 3- and 5- regular graphs have maxima located in the relative vicinity of one another. Subgraphs of 4-regular graphs also have closely positioned maxima between themselves; however, only half of them match with the maxima of subgraphs of odd-regular random graphs.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g002.tif"/>
</fig>
<p>The same effect is observed for subgraphs of 4-regular; see <xref ref-type="fig" rid="F2">Figure 2</xref> (middle row). The optimized parameters are mutually transferable between all four possible subgraphs of 4-regular graphs. Notice, however, that the locations of exactly half of all maxima for the subgraphs of 4-regular graphs do not match with those for 3-regular graphs. This means that one cannot expect good transferability of optimized parameters across MaxCut QAOA instances for 3- and 4-regular random graphs if these optimal parameters are to be transferred directly. It has been recently shown in (<xref ref-type="bibr" rid="B5">Basso et al., 2022</xref>) that gamma parameters can be rescaled in order to generalize between different random d-regular graphs.</p>
<p>Focusing now on all five possible subgraphs of 5-regular graphs, <xref ref-type="fig" rid="F2">Figure 2</xref> (bottom row), we notice that, again, good parameter transferability is expected between all instances of 5-regular random graphs. Moreover, the locations of the maxima match well with those for 3-regular graphs, indicating good transferability across 3- and 5-regular random graphs.</p>
<p>We discuss parameter concentration for instances of random graphs in a later section; similar discussions can be found in (<xref ref-type="bibr" rid="B7">Brandao et al., 2018</xref>) and (<xref ref-type="bibr" rid="B30">Wurtz and Lykov, 2021</xref>) in the context of 3-regular graphs.</p>
<p>To further investigate transferability among regular graphs, we evaluate the subgraph transferability map between all possible subgraphs of <italic>d</italic>-regular graphs, <italic>d</italic> &#x2264; 8; see <xref ref-type="fig" rid="F3">Figure 3</xref>. The top panel shows the colormap of parameter transferability coefficients between all possible pairs of subgraphs of <italic>d</italic>-regular graphs (<italic>d</italic> &#x2264; 8, 35 subgraphs total). Each axis is split into groups of <italic>d</italic> subgraphs of <italic>d</italic>-regular graphs, and the color values in each cell represent the transferability coefficient <italic>T</italic>(<italic>D</italic>, <italic>A</italic>) computed for the corresponding directional pair of subgraphs <italic>D</italic>, <italic>A</italic>, defined as follows. For every subgraph <italic>G</italic>, we performed numerical optimization with 200 steps, repeated 20 times with random initial points. This process results in 20 sets of optimal parameters of the form <inline-formula id="inf5">
<mml:math id="m9">
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> pairs, the best of which we will denote as (<italic>&#x3b3;</italic>
<sub>
<italic>G</italic>&#x2a;</sub>, <italic>&#x3b2;</italic>
<sub>
<italic>G</italic>&#x2a;</sub>). Doing so for the donor subgraph <italic>D</italic> and the acceptor subgraph <italic>A</italic>, the transferability coefficient <italic>T</italic>(<italic>D</italic>, <italic>A</italic>) averages over the QAOA energy contribution of each <inline-formula id="inf6">
<mml:math id="m10">
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> on the acceptor subgraph <italic>A</italic> as follows:<disp-formula id="e3">
<mml:math id="m11">
<mml:mtext>T</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mtext>D</mml:mtext>
<mml:mo>,</mml:mo>
<mml:mtext>A</mml:mtext>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>20</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover accentunder="false" accent="true">
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>20</mml:mn>
</mml:mrow>
</mml:munderover>
</mml:mstyle>
<mml:mfrac>
<mml:mrow>
<mml:mtext>A</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mtext>A</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mo>&#x2a;</mml:mo>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mo>&#x2a;</mml:mo>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:math>
<label>(3)</label>
</disp-formula>where <italic>A</italic>(<italic>&#x3b3;</italic>, <italic>&#x3b2;</italic>) is the QAOA MaxCut energy of subgraph <italic>A</italic> as a function of parameters (<italic>&#x3b3;</italic>, <italic>&#x3b2;</italic>).</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Transferability map between all subgraphs of random regular graphs with maximum node degree <italic>d</italic>
<sub>max</sub> &#x3d; 8, for QAOA depth <italic>p</italic> &#x3d; 1. High (blue) and low (red) values represent good and bad transferability, respectively. Good transferability among even-regular and odd-regular random graphs and poor transferability across even- and odd-regular graphs in both directions are observed.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g003.tif"/>
</fig>
<p>Instead of averaging over the 20 optimal parameters of the donor subgraph, we could have considered only the contribution of the donor&#x2019;s best optimal parameters (<italic>&#x3b3;</italic>
<sub>
<italic>D</italic>&#x2a;</sub>, <italic>&#x3b2;</italic>
<sub>
<italic>D</italic>&#x2a;</sub>) in the above equation. For most donors, however, these best parameters were universal and hence yielded high transferability to most acceptors. However, in practice, because of a lack of iterations or multistarts, we may converge to non-universal optimal parameters, resulting in the donor&#x2019;s poor transferability with some acceptors. The likelihood of converging to these non-universal optima for random graphs is discussed in <xref ref-type="sec" rid="s10">Supplementary Section SA</xref>. Universal and nonuniversal optimal parameters are discussed in detail in <xref ref-type="sec" rid="s4">Section 4</xref>.</p>
<p>This inconsistency was discussed for 3-regular and 4-regular subgraphs earlier in this section. For example, half of the local optima of 3 regular subgraphs have good transferability to 4-regular subgraphs while the half yield poor transferability, as shown in <xref ref-type="fig" rid="F2">Figure 2</xref>. Thus, to reflect practical considerations and avoid such inconsistency, we average over the contributions of 20 optimal parameters of the donor subgraph in Eq. <xref ref-type="disp-formula" rid="e3">3</xref>. It is worth noting here that this averaging over 20 optimal parameters can result in poor transferability, as seen for donor subgraph &#x23;0 to acceptor subgraphs &#x23;2, &#x23;9, &#x23;20, and &#x23;35. For these cases, there is a considerable subset of the donor&#x2019;s optimal parameters that lead to poor transferability. All considered subgraphs are shown in the bottom panel of <xref ref-type="fig" rid="F3">Figure 3</xref>. Note that parameter transferability is a directional property between (sub) graphs, and good transferability from (sub) graph D to (sub) graph A does not guarantee good transferability from A to D. This general fact can be easily understood by considering two graphs with commensurate energy landscapes, for which every energy maximum corresponding to graph D also falls onto the energy maximum for graph A, but some of the energy maxima for graph D do not coincide with those of graph A.</p>
<p>The regular pattern of alternating clusters of high- and low-transferability coefficients in <xref ref-type="fig" rid="F3">Figure 3</xref> illustrates that the parameter transferability effect extends from 3-regular graphs to the entire family of odd-regular graphs, as well as to even-regular graphs, with poor transferability between the two classes. For example, the established result for 3-regular graphs is reflected at the intersection of columns and rows with the label &#x201c;(3)&#x201d; for both donor and acceptor subgraphs. The fact that all cells in the 3 &#xd7; 3 block in <xref ref-type="fig" rid="F3">Figure 3</xref>, corresponding to parameter transfer between subgraphs of 3-regular graphs, have high values, representing high mutual transferability, gives a good indication of optimal QAOA parameter transferability between arbitrary 3-regular graphs (<xref ref-type="bibr" rid="B7">Brandao et al., 2018</xref>).</p>
</sec>
<sec id="s3-2">
<title>3.2 General random graph transferability</title>
<p>Having considered optimal MaxCut QAOA parameter transferability between random regular graphs, we now focus on general random graphs. Subgraphs of an arbitrary random graph differ from subgraphs of random regular graphs in that the two nodes connected by the central edge can have a different number of connected edges, making the set of subgraphs of general random graphs much more diverse. The upper panel of <xref ref-type="fig" rid="F4">Figure 4</xref> shows the transferability map between all possible subgraphs of random graphs with node degrees <italic>d</italic> &#x2264; 6, a total of 56 subgraphs, presented in the lower panel. The transferability map can serve as a lookup table for determining whether optimized QAOA parameters are transferable between any two graphs.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Transferability map between all subgraphs of random graphs with maximum node degree <italic>d</italic>
<sub>max</sub> &#x3d; 6, for QAOA depth <italic>p</italic> &#x3d; 1. Subgraphs are visually separated by dashed lines into groups of subgraphs with the same degrees of the nodes forming the central edge. Solid black rectangles correspond to optimized parameter transferability between subgraphs of random regular graphs (<xref ref-type="fig" rid="F2">Figure 2</xref>).</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g004.tif"/>
</fig>
<p>
<xref ref-type="fig" rid="F4">Figure 4</xref> reveals another important fact about parameter transferability between subgraphs of general random graphs. Subgraphs labeled as (<italic>i</italic>, <italic>j</italic>), where <italic>i</italic> and <italic>j</italic> represent the degrees of the two central nodes of the subgraph, are in general transferable to any other subgraph (<italic>k</italic>, <italic>l</italic>), provided that all {<italic>i</italic>, <italic>j</italic>, <italic>k</italic>, <italic>l</italic>} are either odd or even. This result is a generalization of the transferability result for odd- and even-regular graphs described above. <xref ref-type="fig" rid="F4">Figure 4</xref>, however, shows that a number of pairs of subgraphs with mixed degrees (not only even or odd) also transfer well to other mixed-degree subgraphs, for example, subgraph <italic>&#x23;</italic>20 (3, 4) &#x2192; subgraph <italic>&#x23;</italic>34 (4, 5). The map of subgraph transferability provides a unique tool for identifying smaller donor subgraphs, the optimized QAOA parameters for which are also nearly optimal parameters for the original graph. The map can also be used to define the likelihood of parameter transferability between two graphs based on their subgraphs. As was the case for random regular graphs (see <xref ref-type="fig" rid="F2">Figure 2</xref>), we see clustering of optimal parameters for subgraphs of random graphs in <xref ref-type="fig" rid="F5">Figure 5</xref>.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Distribution of optimal parameters of subgraphs with node degrees of central nodes ranging from 1 to 6 (total 56). Each subgraph was optimized with 20 multistarts, each of which is plotted in the figure above.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g005.tif"/>
</fig>
</sec>
<sec id="s3-3">
<title>3.3 Parameter transferability examples</title>
<p>We will now demonstrate that the parameter transferability map from <xref ref-type="fig" rid="F4">Figure 4</xref> can be used to find small-<italic>N</italic> donor graphs from which the optimized QAOA parameters can be successfully transferred to a MaxCut QAOA instance on a much larger acceptor graph. Initially, we consider three 256-node acceptor graphs to be solved and three 6-node donor graphs; see <xref ref-type="fig" rid="F6">Figure 6</xref>. <xref ref-type="table" rid="T1">Table 1</xref> contains the details of the donor and acceptor graphs, including the total number of edges, their optimized QAOA energies, the energy of the optimal classical solution, and the approximation ratio. Graphs 1 and 4 consist exclusively of odd-degree nodes, graphs 2 and 5 contain roughly the same amount of both odd- and even-degree nodes, and graphs 3 and 6 contain exclusively even-degree nodes. The optimized QAOA parameters for the donor and acceptor graphs were found by performing numerical optimization with 20 restarts, and 200 iterations each. Additionally, we use a greedy ordering algorithm and an RMSprop optimizer, with a learning rate of 0.002. <xref ref-type="table" rid="T2">Table 2</xref> shows the results of the corresponding transfer of optimized parameters from the donor graphs &#x23;&#x23;1&#x2013;3 to the acceptor graphs &#x23;&#x23;4&#x2013;6, correspondingly. The approximation ratios obtained as a result of the parameter transfer in all three cases show only a 1%&#x2013;2% decrease compared with those obtained by optimizing the QAOA parameters for the corresponding acceptor graphs directly. These examples demonstrate the power of the approach introduced in this paper.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Demonstration of optimized parameter transferability between <italic>N</italic> &#x3d; 6 donor and <italic>N</italic> &#x3d; 256 acceptor random graphs. Using optimized parameters from the donor graph for the acceptor leads to the reduction in approximation ratio of 1.0%, 2.6%, and 1.0% for the three examples (top to bottom, compared with optimizing the parameters for the acceptor graph directly, for <italic>p</italic> &#x3d; 1).</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g006.tif"/>
</fig>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Details of donor and acceptor graphs, including number of nodes, number of edges, and both QAOA, and classically optimized energies, along with their corresponding approximation ratios.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Graph</th>
<th align="center">Nodes</th>
<th align="center">Edges</th>
<th align="center">QAOA energy</th>
<th align="center">Energy (Opt)</th>
<th align="center">Approx. Ratio</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">&#x23;1</td>
<td align="center">6</td>
<td align="center">7</td>
<td align="center">4.6481</td>
<td align="center">6.0</td>
<td align="center">0.7746</td>
</tr>
<tr>
<td align="center">&#x23;2</td>
<td align="center">6</td>
<td align="center">6</td>
<td align="center">4.1272</td>
<td align="center">5.0</td>
<td align="center">0.8254</td>
</tr>
<tr>
<td align="center">&#x23;3</td>
<td align="center">6</td>
<td align="center">9</td>
<td align="center">5.7050</td>
<td align="center">6.0</td>
<td align="center">0.9508</td>
</tr>
<tr>
<td align="center">&#x23;4</td>
<td align="center">256</td>
<td align="center">405</td>
<td align="center">269.1192</td>
<td align="center">363.0</td>
<td align="center">0.7413</td>
</tr>
<tr>
<td align="center">&#x23;5</td>
<td align="center">256</td>
<td align="center">461</td>
<td align="center">301.7699</td>
<td align="center">400.0</td>
<td align="center">0.7544</td>
</tr>
<tr>
<td align="center">&#x23;6</td>
<td align="center">256</td>
<td align="center">502</td>
<td align="center">327.4132</td>
<td align="center">430.0</td>
<td align="center">0.7614</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>QAOA, energies from transferred optimal parameters from 6-node donor graphs to 256-node acceptor graphs, along with their corresponding approximation ratios. The values in parenthesis show the reduction in the approximation ratio.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Transfer</th>
<th align="center">QAOA energy</th>
<th align="center">Approx. Ratio</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">&#xa0;&#x23;1 &#x2192; &#x23;4</td>
<td align="center">226.2350</td>
<td align="center">0.7334 (&#x2212;1.0%)</td>
</tr>
<tr>
<td align="center">&#xa0;&#x23;2 &#x2192; &#x23;5</td>
<td align="center">293.8988</td>
<td align="center">0.7347 (&#x2212;2.6%)</td>
</tr>
<tr>
<td align="center">&#xa0;&#x23;3 &#x2192; &#x23;6</td>
<td align="center">323.8726</td>
<td align="center">0.7753 (&#x2212;1.0%)</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>To extend our analysis of parameter transferability between QAOA instances, we perform transferability of optimal parameters between large sets of small donor graphs to a fixed, larger acceptor graph. In particular, we transfer optimal parameters from donors ranging from 6 to 20 nodes to 64-, 128-, and 256-node acceptor graphs. <xref ref-type="fig" rid="F7">Figure 7</xref> shows the approximation ratio as we increase the number of donor graph nodes. These donor graphs were generated starting with graphs of exclusively odd-degree nodes and sequentially increasing the number of even-degree nodes until graphs of exclusively even-degree nodes were obtained. For each increasing number of node in a graph, 100 donor graphs were generated and each of their 20 sets of optimal parameters (20 multistarts) were transferred to the acceptor graph. We see that there are a few cases for which we achieve an approximation ratio that is comparable to the native approximation ratio for each of the acceptor graphs. Most notably, we can achieve good transferability of optimal parameters to larger (i.e., 256-node) acceptor graphs without having to increase the size of our donor graph. Each row of <xref ref-type="fig" rid="F7">Figure 7</xref> corresponds to an increasing acceptor graph size, while each column corresponds to the parity of the acceptor graph (a formal definition and study of parity follow in the next section), with a transition from odd to even parity in graphs going from left to right. For the fully odd and fully even acceptor graphs, we notice a bimodal distribution in approximation ratios. Remarkably, for even acceptor cases, the bimodal distribution has one mode centered around the mean (white dot) and one above the mean. This points to the fact that, regardless of donor graph parity, one can achieve better parameter transferability when transferring optimal parameters to acceptor graphs with even parity. We see this transition from odd to even acceptor graphs in the way the bimodal distribution shifts, there being a monomodal distribution for the cases where the acceptor graphs are neither even nor odd.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Approximation ratios for QAOA parameter transferability between lists of 6- to 20-node donor graphs and <bold>(A&#x2013;C)</bold> 64-node acceptor graphs, <bold>(D&#x2013;F)</bold> 128-node acceptors, and <bold>(G&#x2013;I)</bold> 256-node acceptors. The 64-node acceptors (top row) have the following parities: <bold>(A)</bold> 0.00, <bold>(B)</bold> 0.53, and <bold>(C)</bold> 1.00 eveness; 128-node acceptors (middle row) have the following parities: <bold>(D)</bold> 0.093, <bold>(E)</bold> 0.5, and <bold>(F)</bold> 0.98 evenness; and, 256-node acceptors (bottom row) have the following parities: <bold>(G)</bold> 0, <bold>(H)</bold> 0.49, and <bold>(I)</bold> 1.0 eveness.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g007.tif"/>
</fig>
<p>The reason for this increased likeliness of good transferability to even acceptor graphs will be explored in future work. For now, we turn our focus to parity in graphs as an alternative metric for determining good transferability between donor-acceptor graph pairs, one that does not involve subgraph decomposition (and parameter transferability between individual subgraphs).</p>
</sec>
<sec id="s3-4">
<title>3.4 Parity and transferability</title>
<p>As mentioned previously, the transferability maps of regular and random subgraphs suggest that the parity of graph pairs may affect their transferability. Here, we define parity of a graph <italic>G</italic> &#x3d; (<italic>V</italic>, <italic>E</italic>) to be the proportion of nodes of <italic>G</italic> with an even degree:<disp-formula id="e4">
<mml:math id="m12">
<mml:msub>
<mml:mrow>
<mml:mi>&#x3c0;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="italic">even</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mi>V</mml:mi>
<mml:mo stretchy="false">&#x7c;</mml:mo>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:math>
<label>(4)</label>
</disp-formula>where <italic>n</italic>
<sub>
<italic>even</italic>
</sub> is the number of even nodes in graph <italic>G</italic>. For this and upcoming sections, we focus on transferability between 20-node random graphs. That is, we perform optimal parameter transferability between 20-node donor and 20-node acceptor graphs. For every possible number of even-degree nodes (0, 2, 4, &#x2026; , 20), we generated 10 graphs with distinct degree sequences, resulting in a total of 110 20-node random graphs, with maximum node degree restricted to 6.</p>
<p>The computed transferability coefficients among each graph pair, sorted by their parity, are shown in <xref ref-type="fig" rid="F8">Figure 8</xref>. Each block in the heatmap represents the average transferability coefficient of 100 graph pairs constructed from 10 distinct donor graphs and 10 distinct acceptor graphs. We can see that even graphs, those with <italic>&#x3c0;</italic>
<sub>
<italic>G</italic>
</sub> &#x3d; 0.8&#x2013;1, and odd graphs, those with <italic>&#x3c0;</italic>
<sub>
<italic>G</italic>
</sub> &#x3d; 0&#x2013;0.2, transfer well among themselves. However, the transferability between even donors and odd acceptors, as well as between odd donors and even acceptors, is poor.</p>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>Transferability between 20-node random graphs as a function of the parity of degree of their vertices. The color of each block represents the average transferability of 100 graph pairs. As shown, graph pairs consisting of graphs of similar parity transfer well, while those of different parity transfer poorly.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g008.tif"/>
</fig>
<p>This heatmap also suggests that the mutual transferability of a donor graph is not necessary for its good transferability with other random graphs, where <italic>mutual transferability</italic> of a graph <italic>G</italic> is a measure of how well the subgraphs of <italic>G</italic> transfer among themselves. Formally, it is defined as<disp-formula id="e5">
<mml:math id="m13">
<mml:mtable class="aligned">
<mml:mtr>
<mml:mtd columnalign="right">
<mml:mtext>MT</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>d</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>a</mml:mi>
<mml:mo>&#x2260;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:msub>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>a</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mfrac>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>d</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>a</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>T</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mtd>
</mml:mtr>
</mml:mtable>
</mml:math>
<label>(5)</label>
</disp-formula>where {<italic>G</italic>} is the set of distinct subgraphs of graph <italic>G</italic>, <italic>n</italic>
<sub>
<italic>G</italic>
</sub>(<italic>i</italic>) is the number of edges in <italic>G</italic> having subgraph <italic>i</italic>, and<disp-formula id="equ3">
<mml:math id="m14">
<mml:msub>
<mml:mrow>
<mml:mi>T</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>d</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>a</mml:mi>
<mml:mo>&#x2260;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x22c5;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>G</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>a</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:math>
</disp-formula>is the total number of subgraph pairs consisting of distinct subgraphs within <italic>G</italic>. Graphs with low mutual transferability are those whose subgraphs transfer poorly among themselves. This is true for graphs with a nearly equal number of odd-parity and even-parity subgraphs since subgraph pairs of different parity report poor transferability coefficients, as shown in <xref ref-type="fig" rid="F2">Figures 2</xref>, <xref ref-type="fig" rid="F3">3</xref>. In our case, such graphs are likely to be mixed-parity graphs, in other words, those with <italic>&#x3c0;</italic>
<sub>
<italic>G</italic>
</sub> &#x3d; 0.4&#x2013;0.6. However, the results in <xref ref-type="fig" rid="F8">Figure 8</xref> show that these graphs have good transferability to nearly all random graphs in the data set.</p>
<p>This trend can be explained by analyzing the energy landscapes of subgraphs. Most even- and odd-regular subgraphs have 4 maxima, two of which are universal for all regular subgraphs, as discussed in <xref ref-type="sec" rid="s3-1">Section 3.1</xref>; the same trends were also observed in random subgraphs (see <xref ref-type="fig" rid="F9">Figure 9</xref>). Since energy landscapes of random graphs are sums of the energy landscape of its subgraphs, most 20-node random graphs share the same points, centers 1,2 in <xref ref-type="fig" rid="F10">Figure 10</xref>, as their local or global optima, as shown in <xref ref-type="fig" rid="F11">Figure 11</xref>. On the other hand, the remaining two nonuniversal optimal parameters of regular subgraphs are shared only across regular subgraphs of similar parity. This property is also emergent in random graphs. In <xref ref-type="fig" rid="F11">Figure 11</xref>, only odd random graphs share centers 3, 4 as their local optima, while only even graphs share centers 5, 6 as their local optima. However, mixed parity contained a nearly equal number of odd and even subgraphs. Since nonuniversal maxima of even subgraphs are minima for odd subgraphs and <italic>vice versa</italic>, these nonuniversal local optima blur on the energy landscapes of mixed-parity graphs. As a result, these graphs&#x2019; landscapes contain only universal maxima, as shown in the fourth energy landscape in <xref ref-type="fig" rid="F9">Figure 9</xref>. With only universal parameters as their optimal parameters, mixed-parity graphs should indeed transfer well to all random graphs, as shown in the middle columns of <xref ref-type="fig" rid="F8">Figure 8</xref>.</p>
<fig id="F9" position="float">
<label>FIGURE 9</label>
<caption>
<p>Energy landscapes of some 20-node graphs sorted by parity. Each subplot is the average energy landscape of 10 20-node random graphs with the specified parity.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g009.tif"/>
</fig>
<fig id="F10" position="float">
<label>FIGURE 10</label>
<caption>
<p>Energy landscapes of most 20-node random graphs had either local minima or maxima at one of these 6 centers. Here we label those points for later reference in the text.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g010.tif"/>
</fig>
<fig id="F11" position="float">
<label>FIGURE 11</label>
<caption>
<p>Approximation ratios of the 110 20-node graphs at the 6 points in parameter space identified in <xref ref-type="fig" rid="F9">Figure 9</xref>. Parity, or the number of even-degree nodes in a graph, affects which centers correspond with minima and maxima. The first two centers, however, maximize every graph in the data set.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g011.tif"/>
</fig>
<p>The distribution of optimal parameters also explains poor transferability across random graphs of different parity. In <xref ref-type="fig" rid="F11">Figure 11</xref>, the nonuniversal optimal parameters that maximize the MaxCut energy of odd random graphs, centers 3, 4 also minimize that of even random graphs. Similarly, the nonuniversal optimal parameters that maximize the MaxCut energy of even random graphs, centers 3, 4 also minimize that of odd random graphs. Consequently, transferring nonuniversal optimal parameters of even random graphs to odd random graphs and <italic>vice versa</italic> would result in poor approximation ratios, as evident in <xref ref-type="fig" rid="F8">Figure 8</xref>.</p>
<p>Furthermore, above-average transferability for all graph pairs can be attributed to universal parameters. As shown in <xref ref-type="fig" rid="F12">Figure 12</xref>, all graph pairs have a true similarity or transferability coefficient greater than 0.60. Such a high lower bound can be attributed to universal parameters. Going back to Eq. <xref ref-type="disp-formula" rid="e3">3</xref>, good transferability depends on whether the donor&#x2019;s optimal parameters <inline-formula id="inf7">
<mml:math id="m15">
<mml:mrow>
<mml:mo stretchy="false">(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b3;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>&#x3b2;</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="false">)</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> optimize the acceptor graph as well. If most of the donor&#x2019;s optimal parameters are universal, in other words, are in the vicinity of centers 1, 2 in <xref ref-type="fig" rid="F10">Figure 10</xref>, then the transferability coefficient will be high, regardless of the acceptor graph. In fact, <xref ref-type="sec" rid="s10">Supplementary Figure S2</xref> shows that, on average, all graphs reported at least half of their 20 optimal parameters as universal. As a result, they transfer well to other random graphs.</p>
<fig id="F12" position="float">
<label>FIGURE 12</label>
<caption>
<p>Comparison of subgraph similarity metric <italic>SS</italic> with true similarity for 110<sup>2</sup> graph pairs consisting of 20-node graphs. The color indicates the density of points. For most graph pairs, <italic>SS</italic> underestimates the true similarity.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g012.tif"/>
</fig>
</sec>
<sec id="s3-5">
<title>3.5 Predicting transferability using subgraphs</title>
<p>We have used the transferability coefficient to test whether an acceptor shares the same optimal parameters as its donor. In practice, this quantity is unknown because it requires knowledge of the acceptor&#x2019;s maximum energy. In earlier examples, we used the parity of graphs to explain transferability among random graphs, but the parity of a graph is just one emergent property from its subgraphs. Using subgraphs directly, we devise a subgraph similarity metric <italic>SS</italic> to predict the transferability ratio between a donor graph <italic>D</italic> &#x3d; (<italic>V</italic>
<sub>
<italic>D</italic>
</sub>, <italic>E</italic>
<sub>
<italic>D</italic>
</sub>) and an acceptor graph <italic>A</italic> &#x3d; (<italic>V</italic>
<sub>
<italic>A</italic>
</sub>, <italic>E</italic>
<sub>
<italic>A</italic>
</sub>) as follows:<disp-formula id="e6">
<mml:math id="m16">
<mml:mtext>SS</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>D</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>A</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>d</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mrow>
<mml:mo>&#x2211;</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mi>a</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mfenced open="{" close="}">
<mml:mrow>
<mml:mi>A</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:msub>
<mml:mrow>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>A</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>a</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mfrac>
<mml:mrow>
<mml:mtext>T</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>d</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>a</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>E</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mo>&#x22c5;</mml:mo>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:msub>
<mml:mrow>
<mml:mi>E</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>A</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo stretchy="false">&#x7c;</mml:mo>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:math>
<label>(6)</label>
</disp-formula>where {<italic>G</italic>} is the set of distinct <italic>p</italic> &#x3d; 1 subgraphs of <italic>G</italic>, <italic>n</italic>
<sub>
<italic>G</italic>
</sub>(<italic>g</italic>) is the number of edges in <italic>G</italic> that share the subgraph <italic>g</italic>, and &#x7c;<italic>E</italic>
<sub>
<italic>D</italic>
</sub>&#x7c;&#x22c5;&#x7c;<italic>E</italic>
<sub>
<italic>A</italic>
</sub>&#x7c; is the total number of subgraph pairs across graphs <italic>D</italic> and <italic>A</italic>. Hence, this similarity metric states that the transferability coefficient of a donor and acceptor is the average transferability coefficient of donor subgraph-acceptor subgraph pairs.</p>
<p>In <xref ref-type="fig" rid="F12">Figure 12</xref> we compare this similarity metric with the true similarity or transferability coefficient. While this result does reveal a linear correlation between the two quantities, the metric clearly underapproximates the transferability coefficient by 0.05&#xa0;units on average. In fact, <xref ref-type="fig" rid="F13">Figure 13</xref> shows that graph pairs with mixed-parity graphs as donors report the highest inconsistencies. This poor performance results from their constituent subgraphs. As discussed in <xref ref-type="sec" rid="s3-4">Section 3.4</xref>, mixed-parity graphs consist of a nearly equal number of odd and even subgraphs. When optimized, these donor subgraphs may have nonuniversal optimal parameters. When transferred to an acceptor subgraph, the resulting transferability coefficient may be either poor or good, depending on the parity of that acceptor subgraph. While these nonuniversal optima do affect the subgraph similarity metric <italic>SS</italic>, they do not affect true similarity. As shown in <xref ref-type="fig" rid="F9">Figure 9</xref>, mixed-parity graphs&#x2019; optimal parameters are universal. Thus, they transfer well to any random graph, regardless of its parity. Therefore, the subgraph similarity metric underestimates true similarity because it fails to capture that most optimal parameters of mixed-parity graphs are universal.</p>
<fig id="F13" position="float">
<label>FIGURE 13</label>
<caption>
<p>Differences between subgraph similarity metric <italic>SS</italic> and true similarity sorted by parity of donors and acceptors. <italic>SS</italic> largely underestimates true similarity for graph pairs consisting of graphs of the same parity.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g013.tif"/>
</fig>
</sec>
<sec id="s3-6">
<title>3.6 Predicting transferability using parity</title>
<p>Another approach to predicting transferability or similarity between two graphs is using their parity. In <xref ref-type="sec" rid="s3-4">Section 3.4</xref> we observed that two graphs of similar parity have a high transferability ratio. If this correlation was ideal, then results shown in <xref ref-type="fig" rid="F8">Figure 8</xref> would resemble those in <xref ref-type="fig" rid="F14">Figure 14</xref>. The parity similarity metric <italic>PS</italic> corresponding to the latter figure is easy to compute:<disp-formula id="e7">
<mml:math id="m17">
<mml:mtext>PS</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>D</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>A</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>0.29</mml:mn>
<mml:mo>&#x22c5;</mml:mo>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mtext>Parity</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>D</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo>&#x2212;</mml:mo>
<mml:mtext>Parity</mml:mtext>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>A</mml:mi>
</mml:mrow>
</mml:mfenced>
<mml:mo stretchy="false">&#x7c;</mml:mo>
<mml:mo>.</mml:mo>
</mml:math>
<label>(7)</label>
</disp-formula>
</p>
<fig id="F14" position="float">
<label>FIGURE 14</label>
<caption>
<p>Sorting of parity similarity metric <italic>PS</italic> for graph pairs based on the parity of the donor and the acceptor. Since <italic>PS</italic> assumes that graphs of similar parity transfer well, the diagonal reports the highest <italic>PS</italic>.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g014.tif"/>
</fig>
<p>Thus, this metric penalizes graph pairs consisting of different parity graph pairs. Note that the lowest value of this metric is <inline-formula id="inf8">
<mml:math id="m18">
<mml:mo>&#x2248;</mml:mo>
<mml:mn>0.71</mml:mn>
</mml:math>
</inline-formula>, which is consistent with results from <xref ref-type="fig" rid="F8">Figure 8</xref>. <xref ref-type="fig" rid="F15">Figure 15</xref> illustrates the performance of this new similarity metric. The plot contains discrete columns because it is not possible to generate 20 node graphs with an arbitrary number of even degree nodes. To ensure that the sum of a degree sequence is even, the associated graphs only vary in even degree nodes in increments of two, resulting in parity of 0.0, 0.1, 0.2, &#x2026; 1.0. This discretization is also reflected in the similarity metric.</p>
<fig id="F15" position="float">
<label>FIGURE 15</label>
<caption>
<p>Comparison of parity similarity metric <italic>PS</italic> with true similarity for 110<sup>2</sup> graph pairs consisting of 20-node graphs. The discrete columns occur because we cannot generate 20-node graphs with an arbitrary number of even-degree nodes.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g015.tif"/>
</fig>
<p>To test our similarity metric for the data set shown in <xref ref-type="fig" rid="F7">Figure 7</xref>, we compare our metric with the approximation ratio. <xref ref-type="fig" rid="F16">Figure 16</xref> shows that as parity between donor-acceptor pairs approaches 1 (i.e., the donor and acceptor graphs have the same parity), we achieve a higher approximation ratio. Noticeably, we see that we can have a good approximation ratio even if our parity similarity does not dictate so. This can be attributed to the fact that we are exploiting only one structural feature from our graphs.</p>
<fig id="F16" position="float">
<label>FIGURE 16</label>
<caption>
<p>For an increasing number of donor graph nodes, we see that parity can determine good transferability. For the case with subgraph transferability, we see that this does not depend on the number of nodes of the donor graph.</p>
</caption>
<graphic xlink:href="frqst-02-1200975-g016.tif"/>
</fig>
<p>These results indicate that one can use a parity approach to determine good transferability between donor-acceptor pairs. Furthermore, one can generate a parity metric that caters to specific graphs (please refer to <xref ref-type="sec" rid="s10">Supplementary Material</xref>).</p>
</sec>
</sec>
<sec id="s4">
<title>4 Conclusion and outlook</title>
<p>Finding optimal QAOA parameters is a critical step in solving combinatorial optimization problems by using the QAOA approach. Several existing techniques to accelerate the parameter search are based on advanced optimization and machine learning strategies. In most works, however, various types of global optimizers are employed. Such a straightforward approach is highly inefficient for exploration because of the complex energy landscapes for hard optimization instances.</p>
<p>An alternative effective technique presented in this paper is based on two intuitive observations: 1) The energy landscapes of small subgraphs exhibit &#x201c;well-defined&#x201d; areas of extrema that are not anticipated to be an obstacle for optimization solvers (see <xref ref-type="fig" rid="F2">Figure 2</xref>), and 2) structurally different subgraphs may have similar energy landscapes and optimal parameters. A combination of these observations is important because, in the QAOA approach, the cost is calculated by summing the contributions at the subgraph level, where the size of a subgraph depends on the circuit depth <italic>p</italic>.</p>
<p>With this in mind, the overarching idea of our approach is solving the QAOA parameterization problem for large graphs by optimizing parameterization for much smaller graphs and reusing it. We started with studying the transferability of parameters between all subgraphs of random graphs with a maximum degree of 8. Good transferability of parameters was observed among even-regular and odd-regular subgraphs. At the same time, poor transferability was detected between even- and odd-regular pairs of graphs in both directions, as shown in <xref ref-type="fig" rid="F2">Figures 2</xref>, <xref ref-type="fig" rid="F3">3</xref>. This experimentally confirms the proposed approach.</p>
<p>A remarkable demonstration of random graphs that generalizes the proposed approach is the transferability of the parameters from 6-node random graphs (at the subgraph level) to 256-node random graphs, as shown in <xref ref-type="fig" rid="F6">Figure 6</xref>. The approximation ratio loss of only 1%&#x2013;2% was observed in all three cases. Furthermore, we demonstrated that one need not increase the size of the donor graph to achieve high transferability, even for acceptor graphs with 256 nodes.</p>
<p>Following the subgraph decomposition approach, we showed that one can determine good transferability between donor-acceptor graph pairs by exploiting their similarity based on parity. We see a good correlation between subgraph similarity and parity similarity. In the future, we wish to address the exploitation of graph structure to determine good donor candidates, since subgraph similarities involve overhead calculations of QAOA energies for each pair of donor-acceptor subgraphs.</p>
<p>One may notice that we studied parameter transferability only for <italic>p</italic> &#x3d; 1, where the subgraphs are small and transferability is straightforward. However, our preliminary work suggests that this technique will also work for larger <italic>p</italic>, which will require advanced subgraph exploration algorithms and will be addressed in our following work. In particular, we wish to explore the idea of generating a large database of donor graphs and, together with a graph-embedding technique, obtain optimal QAOA parameters for transferability. We hope that by training a good graph-embedding model, we will be able to apply our technique to various sets of graphs and extend our approach to larger depths. A machine learning approach has been used to determine optimal QAOA parameters (<xref ref-type="bibr" rid="B17">Khairy et al., 2020</xref>), but a study of machine learning for donor graph determination is still an open question.</p>
<p>Another future direction is to determine whether the effects of parity of a graph hold for <italic>p</italic> &#x3e; 1. In particular, we found that the parity of a graph affects the distribution of optimal parameters, as shown in <xref ref-type="fig" rid="F10">Figure 10</xref>; <xref ref-type="sec" rid="s10">Supplementary Figure S2</xref>. It remains to be seen whether parameters concentrate for <italic>p</italic> &#x3e; 1 and, if so, how parity affects their distribution. Analysis of these trends will be critical for the applicability of <italic>PS</italic> for <italic>p</italic> &#x3e; 1.</p>
<p>This work was enabled by the very fast and efficient tensor network simulator QTensor developed at Argonne National Laboratory (<xref ref-type="bibr" rid="B18">Lykov et al., 2021</xref>). Unlike state vector simulators, QTensor can perform energy calculations for most instances with <italic>p</italic> &#x2264; 3, <italic>d</italic> &#x2264; 6 and graphs with <italic>N</italic> &#x223c; 1, 000 nodes very quickly, usually within seconds. For this work we computed QAOA energy for 64-node graphs with <italic>d</italic> &#x2264; 5 at <italic>p</italic> &#x3d; 1, a calculation that took a fraction of a second per each execution on a personal computer. With state vector simulators, however, even such calculations would not have been possible because of the prohibitive memory requirements for storing the state vector.</p>
<p>As a result of this work, finding optimized parameters for some QAOA instances will become quick and efficient, removing this major bottleneck in the QAOA approach and potentially removing the optimization step altogether in some cases, eliminating the variational nature of QAOA. Moreover, our approach will allow finding parameters quickly and efficiently for very large graphs for which it will not be possible to use simulators or other techniques. Our method has important implications for implementing QAOA on relatively slow quantum devices, such as neutral atoms and trapped-ion hardware, for which finding optimal parameters may take a prohibitively long time. Thus, quantum devices will be used only to sample from the output QAOA state to get the final solution to the combinatorial optimization problem. Our work will ultimately bring QAOA one step closer to the realization of quantum advantage.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/<xref ref-type="sec" rid="s10">Supplementary Material</xref>, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec id="s6">
<title>Author contributions</title>
<p>All authors listed have made a substantial, direct, and intellectual contribution to the work and approved it for publication.</p>
</sec>
<sec id="s7">
<title>Funding</title>
<p>This research was developed with funding from the Defense Advanced Research Projects Agency (DARPA). The views, opinions and/or findings expressed are those of the author and should not be interpreted as representing the official views or policies of the Department of Defense or the U.S. Government. AG, DL, XL, IS, JF, and YA are supported in part by funding from the Defense Advanced Research Projects Agency. EG was supported in part by the U.S. Department of Energy, Office of Science, Office of Workforce Development for Teachers and Scientists (WDTS) under the Science Undergraduate Laboratory Internships Program (SULI). This work used in part the resources of the Argonne Leadership Computing Facility, which is a DOE Office of Science User Facility supported under Contract DE-AC02-06CH11357.</p>
</sec>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>Author XL was employed by Fujitsu Research of America, Inc.</p>
<p>The author IS declared that they were an editorial board member of Frontiers, at the time of submission. This had no impact on the peer review process and the final decision.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s9">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="s10">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/frqst.2023.1200975/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/frqst.2023.1200975/full&#x23;supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Image1.TIFF" id="SM1" mimetype="application/TIFF" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image3.TIF" id="SM2" mimetype="application/TIF" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image2.TIF" id="SM3" mimetype="application/TIF" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="DataSheet1.pdf" id="SM4" mimetype="application/pdf" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image5.TIF" id="SM5" mimetype="application/TIF" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image4.TIFF" id="SM6" mimetype="application/TIFF" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<fn-group>
<fn id="fn2">
<label>1</label>
<p>The Gurobi solver provides classically optimal MaxCut solutions in a competitive speed with known optimization gap. For the purpose of this work, there is no particular reason to choose Gurobi over IPOPT or other similarly performing solvers.</p>
</fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Akshay</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Rabinovich</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Campos</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Biamonte</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). <source>Parameter concentration in quantum approximate optimization</source>. <comment>arXiv preprint arXiv:2103.11976</comment>.</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Alexeev</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Bacon</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Brown</surname>
<given-names>K. R.</given-names>
</name>
<name>
<surname>Calderbank</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Carr</surname>
<given-names>L. D.</given-names>
</name>
<name>
<surname>Chong</surname>
<given-names>F. T.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Quantum computer systems for scientific discovery</article-title>. <source>PRX Quantum</source> <volume>2</volume>, <fpage>017001</fpage>. <pub-id pub-id-type="doi">10.1103/prxquantum.2.017001</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Anschuetz</surname>
<given-names>E. R.</given-names>
</name>
<name>
<surname>Kiani</surname>
<given-names>B. T.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Beyond barren plateaus: Quantum variational algorithms are swamped with traps</article-title>. <source>Nat. Commun.</source> <volume>13</volume>, <fpage>7760</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-022-35364-5</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Arute</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Arya</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Babbush</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Bacon</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Bardin</surname>
<given-names>J. C.</given-names>
</name>
<name>
<surname>Barends</surname>
<given-names>R.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Quantum supremacy using a programmable superconducting processor</article-title>. <source>Nature</source> <volume>574</volume>, <fpage>505</fpage>&#x2013;<lpage>510</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-019-1666-5</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Basso</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Farhi</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Marwaha</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Villalonga</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>The quantum approximate optimization algorithm at high depth for MaxCut on large-girth regular graphs and the sherrington-kirkpatrick model</article-title>,&#x201d; in <source>17th conference on the theory of quantum computation, communication and cryptography (TQC 2022)</source>. <source>Leibniz international proceedings in informatics (LIPIcs)</source>. Editors <person-group person-group-type="editor">
<name>
<surname>Le Gall</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Morimae</surname>
<given-names>T.</given-names>
</name>
</person-group> (<publisher-loc>Dagstuhl, Germany</publisher-loc>: <publisher-name>Schloss Dagstuhl &#x2013; Leibniz-Zentrum f&#xfc;r Informatik</publisher-name>), <volume>232</volume>, <fpage>7:1</fpage>&#x2013;<lpage>7:21</lpage>. <pub-id pub-id-type="doi">10.4230/LIPIcs.TQC.2022.7</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Biamonte</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Bergholm</surname>
<given-names>V.</given-names>
</name>
</person-group> (<year>2017</year>). <source>Tensor networks in a nutshell</source>. <comment>[Dataset]</comment>.</citation>
</ref>
<ref id="B7">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Brandao</surname>
<given-names>F. G.</given-names>
</name>
<name>
<surname>Broughton</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Farhi</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Gutmann</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Neven</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2018</year>). <source>For fixed control parameters the quantum approximate optimization algorithm&#x2019;s objective function value concentrates for typical instances</source>. <comment>arXiv preprint arXiv:1812.04170</comment>.</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Egger</surname>
<given-names>D. J.</given-names>
</name>
<name>
<surname>Marecek</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Woerner</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Warm-starting quantum optimization</article-title>. <comment>[Dataset]</comment>
</citation>
</ref>
<ref id="B9">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Farhi</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Goldstone</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Gutmann</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2014</year>). <source>A quantum approximate optimization algorithm</source>. <comment>arXiv preprint arXiv:1411.4028</comment>.</citation>
</ref>
<ref id="B10">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Gurobi Optimization</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2021</year>). <source>Gurobi optimizer reference manual</source>. <comment>[Dataset]</comment>.</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hadfield</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>O&#x2019;Gorman</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Rieffel</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Venturelli</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Biswas</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>From the Quantum Approximate Optimization Algorithm to a quantum alternating operator ansatz</article-title>. <source>Algorithms</source> <volume>12</volume>, <fpage>34</fpage>. <pub-id pub-id-type="doi">10.3390/a12020034</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Herman</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Googin</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Galda</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Safro</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <source>A survey of quantum computing for finance</source>. <comment>arXiv preprint arXiv:2201.02773</comment>.</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hogg</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Portnov</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2000</year>). <article-title>Quantum optimization</article-title>. <source>Inf. Sci.</source> <volume>128</volume>, <fpage>181</fpage>&#x2013;<lpage>197</lpage>. <pub-id pub-id-type="doi">10.1016/s0020-0255(00)00052-9</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hogg</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2000</year>). <article-title>Quantum search heuristics</article-title>. <source>Phys. Rev. A</source> <volume>61</volume>, <fpage>052311</fpage>. <pub-id pub-id-type="doi">10.1103/physreva.61.052311</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Joseph</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Porter</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Castelli</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Geyko</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Graziani</surname>
<given-names>F.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>Quantum computing for fusion energy science applications</article-title>. <source>Phys. Plasmas</source> <volume>30</volume>, <fpage>010501</fpage>. <pub-id pub-id-type="doi">10.1063/5.0123765</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kardashin</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Uvarov</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Biamonte</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Quantum machine learning tensor network states</article-title>. <source>Front. Phys.</source> <volume>8</volume>, <fpage>586374</fpage>. <pub-id pub-id-type="doi">10.3389/fphy.2020.586374</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Khairy</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Shaydulin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Cincio</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Alexeev</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Balaprakash</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2020</year>). &#x201c;<article-title>Learning to optimize variational quantum circuits to solve combinatorial problems</article-title>,&#x201d; in <conf-name>Proceedings of the AAAI Conference on Artificial Intelligence</conf-name>, <fpage>2367</fpage>&#x2013;<lpage>2375</lpage>. <pub-id pub-id-type="doi">10.1609/aaai.v34i03.5616</pub-id>
<volume>34</volume>
</citation>
</ref>
<ref id="B18">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Lykov</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Galda</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Alexeev</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2021</year>). <source>Qtensor</source>. <comment>[Dataset]</comment>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://github.com/danlkv/qtensor">https://github.com/danlkv/qtensor</ext-link>
</comment>.</citation>
</ref>
<ref id="B19">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Lykov</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Schutski</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Galda</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Vinokur</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Alexeev</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2020</year>). <source>Tensor network quantum simulator with step-dependent parallelization</source>. <comment>
<italic>arXiv preprint arXiv:2012.02430</italic>
</comment>.</citation>
</ref>
<ref id="B20">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Newman</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2018</year>). <source>Networks</source>. <publisher-name>Oxford University Press</publisher-name>.</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Outeiral</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Strahm</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Morris</surname>
<given-names>G. M.</given-names>
</name>
<name>
<surname>Benjamin</surname>
<given-names>S. C.</given-names>
</name>
<name>
<surname>Deane</surname>
<given-names>C. M.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>The prospects of quantum computing in computational molecular biology</article-title>. <source>Wiley Interdiscip. Rev. Comput. Mol. Sci.</source> <volume>11</volume>, <fpage>e1481</fpage>. <pub-id pub-id-type="doi">10.1002/wcms.1481</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Preskill</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Quantum computing in the nisq era and beyond</article-title>. <source>Quantum</source> <volume>2</volume>, <fpage>79</fpage>. <pub-id pub-id-type="doi">10.22331/q-2018-08-06-79</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shaydulin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Hadfield</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Hogg</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Safro</surname>
<given-names>I.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Classical symmetries and the quantum approximate optimization algorithm</article-title>. <source>Quantum Inf. Process.</source> <volume>20</volume>, <fpage>359</fpage>. <pub-id pub-id-type="doi">10.1007/s11128-021-03298-4</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Shaydulin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Safro</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Larson</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2019a</year>). &#x201c;<article-title>Multistart methods for quantum approximate optimization</article-title>,&#x201d; in <conf-name>2019 IEEE High Performance Extreme Computing Conference (HPEC)</conf-name> (<publisher-name>IEEE</publisher-name>), <fpage>1</fpage>&#x2013;<lpage>8</lpage>.</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shaydulin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Ushijima-Mwesigwa</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Negre</surname>
<given-names>C. F.</given-names>
</name>
<name>
<surname>Safro</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Mniszewski</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Alexeev</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2019b</year>). <article-title>A hybrid approach for solving optimization problems on small quantum computers</article-title>. <source>Computer</source> <volume>52</volume>, <fpage>18</fpage>&#x2013;<lpage>26</lpage>. <pub-id pub-id-type="doi">10.1109/mc.2019.2908942</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shaydulin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Ushijima-Mwesigwa</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Safro</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Mniszewski</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Alexeev</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2019c</year>). <article-title>Network community detection on small quantum computers</article-title>. <source>Adv. Quantum Technol.</source> <volume>2</volume>, <fpage>1900029</fpage>. <pub-id pub-id-type="doi">10.1002/qute.201900029</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Streif</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Leib</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Training the quantum approximate optimization algorithm without access to a quantum processing unit</article-title>. <source>Quantum Sci. Technol.</source> <volume>5</volume>, <fpage>034008</fpage>. <pub-id pub-id-type="doi">10.1088/2058-9565/ab8c2b</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ushijima-Mwesigwa</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Shaydulin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Negre</surname>
<given-names>C. F.</given-names>
</name>
<name>
<surname>Mniszewski</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Alexeev</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Safro</surname>
<given-names>I.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Multilevel combinatorial optimization across quantum architectures</article-title>. <source>ACM Trans. Quantum Comput.</source> <volume>2</volume>, <fpage>1</fpage>&#x2013;<lpage>29</lpage>. <pub-id pub-id-type="doi">10.1145/3425607</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Fontana</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Cerezo</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Sharma</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Sone</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Cincio</surname>
<given-names>L.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Noise-induced barren plateaus in variational quantum algorithms</article-title>. <source>Nat. Commun.</source> <volume>12</volume>, <fpage>6961</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-021-27045-6</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Wurtz</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Lykov</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2021</year>). <source>The fixed angle conjecture for QAOA on regular MaxCut graphs</source>. <comment>[Dataset]</comment>.</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhou</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.-T.</given-names>
</name>
<name>
<surname>Choi</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Pichler</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Lukin</surname>
<given-names>M. D.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Quantum approximate optimization algorithm: Performance, mechanism, and implementation on near-term devices</article-title>. <source>Phys. Rev. X</source> <volume>10</volume>, <fpage>021067</fpage>. <pub-id pub-id-type="doi">10.1103/PhysRevX.10.021067</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>