<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Mol. Biosci.</journal-id>
<journal-title>Frontiers in Molecular Biosciences</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Mol. Biosci.</abbrev-journal-title>
<issn pub-type="epub">2296-889X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1253689</article-id>
<article-id pub-id-type="doi">10.3389/fmolb.2023.1253689</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Molecular Biosciences</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Exploring rigid-backbone protein docking in biologics discovery: a test using the DARPin scaffold</article-title>
<alt-title alt-title-type="left-running-head">Gaudreault et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fmolb.2023.1253689">10.3389/fmolb.2023.1253689</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Gaudreault</surname>
<given-names>Francis</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1715609/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Baardsnes</surname>
<given-names>Jason</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Martynova</surname>
<given-names>Yuliya</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Dachon</surname>
<given-names>Aurore</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Hogues</surname>
<given-names>Herv&#xe9;</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Corbeil</surname>
<given-names>Christopher R.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1757691/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Purisima</surname>
<given-names>Enrico O.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/881844/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Arbour</surname>
<given-names>M&#xe9;lanie</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2383801/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Sulea</surname>
<given-names>Traian</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1611052/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Human Health Therapeutics Research Centre</institution>, <institution>National Research Council Canada</institution>, <addr-line>Montreal</addr-line>, <addr-line>QC</addr-line>, <country>Canada</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Institute of Parasitology</institution>, <institution>McGill University</institution>, <addr-line>Montreal</addr-line>, <addr-line>QC</addr-line>, <country>Canada</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/175790/overview">F. Javier Luque</ext-link>, University of Barcelona, Spain</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/338042/overview">Pablo Chacon</ext-link>, Spanish National Research Council (CSIC), Spain</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1090431/overview">Baldomero Oliva</ext-link>, Pompeu Fabra University, Spain</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Traian Sulea, <email>traian.sulea@nrc-cnrc.gc.ca</email>
</corresp>
</author-notes>
<pub-date pub-type="epub">
<day>24</day>
<month>08</month>
<year>2023</year>
</pub-date>
<pub-date pub-type="collection">
<year>2023</year>
</pub-date>
<volume>10</volume>
<elocation-id>1253689</elocation-id>
<history>
<date date-type="received">
<day>05</day>
<month>07</month>
<year>2023</year>
</date>
<date date-type="accepted">
<day>14</day>
<month>08</month>
<year>2023</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2023 Gaudreault, Baardsnes, Martynova, Dachon, Hogues, Corbeil, Purisima, Arbour and Sulea.</copyright-statement>
<copyright-year>2023</copyright-year>
<copyright-holder>Gaudreault, Baardsnes, Martynova, Dachon, Hogues, Corbeil, Purisima, Arbour and Sulea</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Accurate protein-protein docking remains challenging, especially for artificial biologics not coevolved naturally against their protein targets, like antibodies and other engineered scaffolds. We previously developed ProPOSE, an exhaustive docker with full atomistic details, which delivers cutting-edge performance by allowing side-chain rearrangements upon docking. However, extensive protein backbone flexibility limits its practical applicability as indicated by unbound docking tests. To explore the usefulness of ProPOSE on systems with limited backbone flexibility, here we tested the engineered scaffold DARPin, which is characterized by its relatively rigid protein backbone. A prospective screening campaign was undertaken, in which sequence-diversified DARPins were docked and ranked against a directed epitope on the target protein BCL-W. In this proof-of-concept study, only a relatively small set of 2,213 diverse DARPin interfaces were selected for docking from the huge theoretical library from mutating 18 amino-acid positions. A computational selection protocol was then applied for enrichment of binders based on normalized computed binding scores and frequency of binding modes against the predefined epitope. The top-ranked 18 designed DARPin interfaces were selected for experimental validation. Three designs exhibited binding affinities to BCL-W in the nanomolar range comparable to control interfaces adopted from known DARPin binders. This result is encouraging for future screening and engineering campaigns of DARPins and possibly other similarly rigid scaffolds against targeted protein epitopes. Method limitations are discussed and directions for future refinements are proposed.</p>
</abstract>
<kwd-group>
<kwd>binding affinity</kwd>
<kwd>protein-protein docking</kwd>
<kwd>rigid backbone</kwd>
<kwd>DARPin</kwd>
<kwd>ProPOSE</kwd>
</kwd-group>
<contract-sponsor id="cn001">Alliance de recherche num&#xe9;rique du Canada<named-content content-type="fundref-id">10.13039/501100021202</named-content>
</contract-sponsor>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Biological Modeling and Simulation</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<title>1 Introduction</title>
<p>Biologics have witnessed a tremendous growth in the past decades, with antibody-based therapeutics leading the way and recombinant proteins forming another important market segment (<xref ref-type="bibr" rid="B8">DeFrancesco, 2019</xref>; <xref ref-type="bibr" rid="B22">Lu et al., 2020</xref>; <xref ref-type="bibr" rid="B17">Kaplon et al., 2023</xref>). Advances in computational methods have spurred the idea that in the not-so-distant future, novel biologics can be discovered entirely <italic>in silico</italic>, complementing current wet-lab methods such as immunization and display technologies. This emerging field is dubbed <italic>de novo</italic> discovery of biologics with a particular emphasis on <italic>de novo</italic> antibody engineering (<xref ref-type="bibr" rid="B10">Fischman and Ofran, 2018</xref>).</p>
<p>Central to this <italic>de novo</italic> discovery approach is the ability to dock and score large libraries of biologic variants on the three-dimensional (3D) structure of a target protein (e.g., the antigen in the case of antibodies). Artificial intelligence/machine learning (AI/ML)-based methods like AlphaFold2 (<xref ref-type="bibr" rid="B16">Jumper et al., 2021</xref>), which have recently demonstrated a tremendous success in predicting protein structures and complexes of biologically co-evolved proteins, unfortunately are not applicable to docking and scoring of antibodies and artificially designed proteins (<xref ref-type="bibr" rid="B39">Yin et al., 2022</xref>). This limitation is due to co-evolution data being essential to AI/ML&#x2019;s success in protein-protein docking (<xref ref-type="bibr" rid="B9">Evans et al., 2022</xref>; <xref ref-type="bibr" rid="B12">Gao et al., 2022</xref>). Compounding the docking and scoring challenge is the difficulty to predict 3D structures of antibody libraries. While there has been some recent success in modeling antibodies with AI/ML methods without co-evolutionary information, there are still challenges in predicting the conformation of the hypervariable CHR-H3 loop (<xref ref-type="bibr" rid="B1">Abanades et al., 2022</xref>; <xref ref-type="bibr" rid="B6">Cohen et al., 2022</xref>; <xref ref-type="bibr" rid="B29">Ruffolo et al., 2022</xref>). Due to technical limitations from the high dimensionality of the CDR-H3 conformational space, the applicability of <italic>de novo</italic> antibody discovery efforts based on docking modeled antibody libraries to an antigen structure was met with limited success, as reported with several classical approaches (<xref ref-type="bibr" rid="B2">Adolf-Bryfogle et al., 2018</xref>; <xref ref-type="bibr" rid="B5">Chowdhury et al., 2018</xref>; <xref ref-type="bibr" rid="B37">Warszawski et al., 2019</xref>; <xref ref-type="bibr" rid="B38">Wood, 2021</xref>). Instead, applications on biologics displaying limited amounts of flexibility should be explored for increased likelihood of success (<xref ref-type="bibr" rid="B40">Youn et al., 2017</xref>; <xref ref-type="bibr" rid="B27">Radom et al., 2019</xref>).</p>
<p>We previously developed ProPOSE, an exhaustive direct protein-protein docker with full atomistic details (<xref ref-type="bibr" rid="B13">Hogues et al., 2018</xref>). By allowing side-chain rearrangements upon docking, ProPOSE delivers the current leading-edge performance in both general protein-protein docking and the specific case of antibody-antigen docking, when the backbone conformations of the interacting partners in the complex are <italic>a priori</italic> known. More specifically, ProPOSE maintains a strong performance even when side-chain flexibility is of concern. However, the docking accuracy was lower when backbone atoms experienced significant displacements between the bound and unbound states. We anticipated that despite its limitations, ProPOSE should be able to show utility in <italic>de novo</italic> biologics discovery when there is limited backbone flexibility upon binding and when reasonable models of backbone conformations can be inferred for the library of potential binders.</p>
<p>Hence, in this proof-of-concept study, we turned away from antibodies and towards the well-known engineered scaffold called DARPin (Designed Ankyrin Repeat Protein) (<xref ref-type="bibr" rid="B4">Binz et al., 2003</xref>). The DARPin scaffold has been refined over the years and has proven its value for the discovery of molecules with various medical and engineering applications, for example, as biotherapeutics, diagnostic agents, biosensors, molecular probes and crystallization helpers (<xref ref-type="bibr" rid="B24">Pluckthun, 2015</xref>; <xref ref-type="bibr" rid="B28">Rothenberger et al., 2022</xref>; <xref ref-type="bibr" rid="B36">Strittmatter et al., 2022</xref>). Compared to antibodies, DARPins are generally considered to be more rigid due to their smaller size and more defined structure. The repeating ankyrin unit (a &#x3b2;-turn followed by two anti-parallel &#x3b1;-helices) confers rigidity and stability to their structure (<xref ref-type="bibr" rid="B18">Kramer et al., 2010</xref>; <xref ref-type="bibr" rid="B30">Schilling et al., 2022</xref>). Such a limited backbone flexibility thus appears suitable for modeling DARPin substitution variants relatively reliably starting from available DARPin template structures.</p>
<p>Hence, the exploratory prospective study described here was centered around applying ProPOSE rigid-backbone docking to the DARPin scaffold exhibiting relative backbone rigidity. A computational flow was devised to generate a relatively small library of diverse DARPin interfaces for directed docking to a known epitope on the structure of the protein target, BCL-W. A selection procedure was further devised to establish a score threshold that captured self-consistent positive controls generated within the same computational procedure. Prospective computational designs were then subjected to experimental testing. Testing of 18 top-ranked hits demonstrated that half of them had detected binding to the target. Comparative analysis of computational and experimental data prompted to several limitations and areas for future improvements of the rigid-docking based approach for <italic>de novo</italic> biologics discovery.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>2 Materials and methods</title>
<sec id="s2-1">
<title>2.1 Computational methods</title>
<p>The sequence-based and structure-based computational design process (<xref ref-type="fig" rid="F1">Figure 1</xref>) consisted of 6 steps which are described in the following sub-sections.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Flowchart of the overall computational design and experimental testing. The first three steps of the computational design are in the sequence space, while the last three steps are in the 3D-structure space and inherit structural knowledge from the Protein Data Bank (<xref ref-type="bibr" rid="B3">Berman et al., 2000</xref>). The main steps of the computational design are numbered outside the boxes and described in the text.</p>
</caption>
<graphic xlink:href="fmolb-10-1253689-g001.tif"/>
</fig>
<sec id="s2-1-1">
<title>2.1.1 Defining the DARPin common framework sequence</title>
<p>Hundreds of DARPin structures with various topologies were published in the literature and are accessible in the PDB, among which many have 4 or 5 repeated ankyrin motifs. Two DARPins evolved through ribosome display to bind BCL-W, and corresponding to PDB entries 4k5a and 4k5b (<xref ref-type="bibr" rid="B32">Schilling et al., 2014b</xref>), were used as known binders in this study. These known binders engage the target in a binding mode which is typical for DARPins, which consists of interactions made by the concave paratope formed by their 5 repeated ankyrin motifs (<xref ref-type="bibr" rid="B4">Binz et al., 2003</xref>; <xref ref-type="bibr" rid="B18">Kramer et al., 2010</xref>; <xref ref-type="bibr" rid="B24">Pluckthun, 2015</xref>; <xref ref-type="bibr" rid="B30">Schilling et al., 2022</xref>). By inspecting the sequences and structures of these known binders and other DARPins with available crystal structures in PDB, a common framework sequence was defined for further library expansion. The main features considered during the selection of a DARPin common framework sequence were: 1) 157 amino acids starting with DLGKK and ending with LQKAA sequences; 2) conserved regions at these N- and C-terminal ends; 3) consensus residues deemed essential for the stability of the overall fold along repeated ankyrin motifs; and 4) key residues contributing to binding along repeated ankyrin motifs. These criteria led to a single DARPin common framework sequence, which corresponded to the DARPin of chain F in the PDB entry 4drx (the nomenclature 4drx [F] is used) (<xref ref-type="bibr" rid="B23">Pecqueur et al., 2012</xref>).</p>
</sec>
<sec id="s2-1-2">
<title>2.1.2 Expanding the framework sequence into a DARPin library</title>
<p>A set of 18 amino-acid positions within the defined DARPin common framework sequence were manually selected and allowed to vary (referred to as variable positions). These positions, which have high-frequency rates of mutation as observed from sequence alignments of many DARPins from the literature, are: 45, 46, 48, 56, 57, 78, 79, 81, 89, 90, 111, 112, 114, 122, 123, 144, 145 and 147 (standard DARPin numbering is applied). Amino-acid side chains at these positions are lining the concave face of the DARPin scaffold by being located within the &#x3b2;-turn loops and following short &#x3b1;-helices of the ankyrin repeats (<xref ref-type="fig" rid="F2">Figure 2</xref>).</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Variable positions on the DARPin scaffold docked onto BCL-W target epitope. <bold>(A)</bold> Sequence alignment between the common framework sequence (4drx [F]), known binders (4k5a [B] and 4k5b [B]), and the positive controls (PC1 and PC2) grafting the interface of the known binders onto the common template sequence, at the 18 variable positions (marked by green Xs). The conventional DARPin sequence numbering scheme is used, <italic>h</italic> denotes &#x3b1;-helix, IR1 to IR3 delineate internal ankyrin repeats 1-3, and N-Cap and C-Cap are the terminal ankyrin repeats. <bold>(B)</bold> Location of the 18 variable positions (spheres) on the 4 DARPin template structures (C&#x3b1;-traces with different shades of green). <bold>(C)</bold> Location of the docking site on the BCL-W target protein indicated by the crystal structure (4k5b) of a known DARPin binder (red cartoon) complexed with the BCL-W target (molecular surface).</p>
</caption>
<graphic xlink:href="fmolb-10-1253689-g002.tif"/>
</fig>
<p>Two positive controls having the 18 variable positions corresponding to known DARPin binders of BCL-W, having PDB entries 4k5a [B] and 4k5b [B], were also built manually into the library. It is important to note that these constructed positive controls share the common framework of the designed library described above and thus differ at several positions from the frameworks of the originating known binders (<xref ref-type="fig" rid="F2">Figure 2</xref>).</p>
</sec>
<sec id="s2-1-3">
<title>2.1.3 Selecting a DARPin sub-library of diverse sequences</title>
<p>An alphabet was created to group amino acids by chemical properties. The following five groups excluding Gly, Cys and Pro were defined: positively-charged (Arg, His, Lys); negatively-charged (Asp, Glu); polar (Asn, Gln, Ser, Thr); non-polar (Ala, Ile, Leu, Met, Val); and aromatic (Phe, Trp, Tyr). Equal probability was given to each group to be selected when mutating sequences. Similarly, amino acids within a group were given equal probability.</p>
<p>The designs were generated using a stochastic procedure in which variable amino-acid positions were mutated either through point mutations or through permutations of amino acids. Multiple starting points in the sequence space were used to generate the designs. The set of mutated designs (M-set) were generated starting from the 4drx [F] sequence chosen as common framework. All variable amino acids were forced to be mutated in this set. To be included in the library, a design sequence had to be sufficiently distant to the designs comprised within the same set. A threshold distance of 10 was fixed which required at least 10 alphabet group changes. The set of permutated designs (P-set) were generated starting from the sequences of the two positive controls. No change in the alphabet group was imposed for this set. A threshold distance of 13 was set, requiring at least 13 amino-acid changes.</p>
</sec>
<sec id="s2-1-4">
<title>2.1.4 Grafting DARPin sequences onto template structures</title>
<p>The designed sub-library sequences were grafted onto four DARPin template structures followed by side-chain repacking using SCWRL4 (<xref ref-type="bibr" rid="B19">Krivov et al., 2009</xref>). The last two alanine residues at the C-terminus of the template sequence were truncated for modeling purposes. Only those side-chains that are different at a given side-chain were mutated and repacked to preserve the structural integrity of the original crystal structures of the DARPin templates. The entire structure was then allowed to be repacked. The DARPin templates from the following PDB entries were used in this study: 4drx [F], 4j7w [A], 5lw2 [A] and 5le6 [A] (<xref ref-type="fig" rid="F1">Figure 1</xref>). The backbone structures of these templates are distinct from those of the two known DARPin binders of BCL-W, 4k5a [B] and 4k5b [B], which were purposely excluded as structural templates to avoid the cognate-docking bias. The selected DARPin template structures underwent the following preparation procedure: 1) addition of missing side chain atoms (no repacking); 2) addition of missing hydrogen atoms and assignment of standard protonation states at pH 7; 3) optimization of the hydrogen-bond network the minH program (<xref ref-type="bibr" rid="B14">Hogues et al., 2014</xref>); and 4) AMBER force-field (<xref ref-type="bibr" rid="B7">Cornell et al., 1995</xref>; <xref ref-type="bibr" rid="B15">Hornak et al., 2006</xref>) energy minimization of added hydrogen atoms and any newly added side-chain atoms with harmonic restraints on all the other heavy atoms of 1,000&#xa0;kcal/mol/ <inline-formula id="inf1">
<mml:math id="m1">
<mml:mrow>
<mml:msup>
<mml:mover accent="true">
<mml:mi mathvariant="normal">A</mml:mi>
<mml:mo>&#x30a;</mml:mo>
</mml:mover>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> followed by energy minimization of the entire structure with harmonic restraints of 10&#xa0;kcal/mol/<inline-formula id="inf2">
<mml:math id="m2">
<mml:mrow>
<mml:msup>
<mml:mover accent="true">
<mml:mi mathvariant="normal">A</mml:mi>
<mml:mo>&#x30a;</mml:mo>
</mml:mover>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> on backbone heavy atoms, 1&#xa0;kcal/mol/<inline-formula id="inf3">
<mml:math id="m3">
<mml:mrow>
<mml:msup>
<mml:mover accent="true">
<mml:mi mathvariant="normal">A</mml:mi>
<mml:mo>&#x30a;</mml:mo>
</mml:mover>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> on side-chain heavy atoms, and no restraints on hydrogen atoms.</p>
</sec>
<sec id="s2-1-5">
<title>2.1.5 DARPin docking protocols</title>
<p>The BCL-W docking-based screening of the DARPin library was performed using the exhaustive docking engine ProPOSE version 1.03 (<xref ref-type="bibr" rid="B13">Hogues et al., 2018</xref>). ProPOSE was run with default parameters using the HITSET flag to force binding towards the set of residues involved in binding BCL-W. Initially, no binding location (or epitope) was defined on the target protein BCL-W and exhaustive docking was performed all around the BCL-W structure. Two BCL-W structures were employed for docking, with PDB entries 4k5a [A] and 4k5b [C] (<xref ref-type="fig" rid="F1">Figure 1</xref>), which correspond to the BCL-W complexed with the two known DARPin binders. For each DARPin library sequence, the four DARPin structural templates carrying the grafted designed sequence were docked against the two BCL-W target structure, resulting in 8 docking experiments. In this study, only the top-1 scored pose generated by ProPOSE was considered for a given complex given its accuracy in pose recovery as top-1 when the protein backbone conformation is known, without the need for rescoring (<xref ref-type="bibr" rid="B13">Hogues et al., 2018</xref>). On average, a single docking run took 30&#xa0;min to execute when parallelized on an Intel Xeon Gold 5,218 using 6 cores.</p>
<p>Epitope restriction on the BCL-W target was introduced after all docking calculations were completed. In this proof-of-concept study, we elected to target the same BCL-W epitope and the DARPin binding mode observed for the two known BCL-W DARPin binders (PDB entries 4k5a and 4k5b) (<xref ref-type="bibr" rid="B32">Schilling et al., 2014b</xref>). The similarity of predicted docked poses of designed sequences relative to these known structures was based on CAPRI classification (<xref ref-type="bibr" rid="B21">Lensink et al., 2017</xref>). Predictions were compared on the basis of: 1) the backbone RMSD of the ligand upon target superposition; 2) the backbone RMSD of the interface upon superposition of interface atoms; and 3) the fraction of preserved contacts (f<sub>con</sub>). The ligand and target were DARPin and BCL-W, respectively. Noteworthy, f<sub>con</sub> was used rather than the standard f<sub>nat</sub> from CAPRI that is derived from the comparison to a native structure. Moreover, f<sub>con</sub> is a position-dependent (amino acid-independent) measure allowing designs with different sequences to be compared. Two predictions were declared as having high, medium or acceptable quality, or as incorrect otherwise, with thresholds defined by the CAPRI classification (<xref ref-type="bibr" rid="B21">Lensink et al., 2017</xref>).</p>
</sec>
<sec id="s2-1-6">
<title>2.1.6 Ranking docked DARPin structures</title>
<p>The number of top-1 scored poses, N<sub>pose</sub>, docked at the targeted epitope from the 8 docking runs was used to retain only those designs that have at least 2 poses docked at the target epitope. To this end, the predicted poses for a given DARPin were grouped using a greedy clustering algorithm with a tolerance of at least medium quality between cluster representatives. In geometric terms, for poses to be considered bound at the targeted epitope occupied by one of the known binders, they were required to have acceptable quality criteria, i.e., 1) f<sub>con</sub> of at least 30% with a ligand backbone RMSD &#x3e;5.0&#xc5; and interface backbone RMSD &#x3e;2.0&#xa0;&#xc5;; or alternatively, 2) f<sub>con</sub> between 10% and 30% while having a ligand backbone RMSD &#x3c;10.0&#xa0;&#xc5; or an interface backbone RMSD &#x3c;4.0&#xa0;&#xc5;. For poses to be part of the same cluster, they were required to have medium quality criteria, i.e., 1) f<sub>con</sub> of at least 50% with ligand backbone RMSD &#x3e;1.0&#xa0;&#xc5; and interface backbone RMSD &#x3e;1.0&#xa0;&#xc5;; or alternatively, 2) f<sub>con</sub> between 30% and 50% while having ligand backbone RMSD &#x3c;5.0&#xa0;&#xc5; or an interface backbone RMSD &#x3c;2.0&#xa0;&#xc5;. No cut-off in score was applied for the clustering.</p>
<p>For each design with N<sub>pose</sub> &#x3e; 1, a consensus score was derived as the arithmetic average over the docking scores of the poses binding to the targeted epitope. Consensus scores over the designs with N<sub>pose</sub> &#x3e; 1 were also normalized into Z-scores to better inform the selection of a top-ranked population based on a minimum number of standard deviations away from the mean calculated from the distribution of all DARPins combining the P-set designs, M-set designs and the positive controls.</p>
</sec>
<sec id="s2-1-7">
<title>2.1.7 Other software and data availability</title>
<p>Structure visualization was performed in PyMOL (The PyMOL Molecular Graphics System, Version 2.0, Schr&#xf6;dinger, LLC). Statistical analyzes were run in R (<xref ref-type="bibr" rid="B26">R Development Core Team, 2011</xref>). ClustalW2 was used to run the multiple sequence alignments (<xref ref-type="bibr" rid="B20">Larkin et al., 2007</xref>).</p>
<p>The sequence datasets generated for this study have been made available as a MongoDB with example scripts that can be found at the GitHub repository <ext-link ext-link-type="uri" xlink:href="https://github.com/gaudreaultfnrc/Darpins">https://github.com/gaudreaultfnrc/Darpins</ext-link>.</p>
</sec>
</sec>
<sec id="s2-2">
<title>2.2 Experimental methods</title>
<sec id="s2-2-1">
<title>2.2.1 Protein expression and purification</title>
<p>Each DARPin design included a N-terminus tag (MRGSHHHHHHGS) and two alanines at their C-terminus as described in (<xref ref-type="bibr" rid="B31">Schilling et al., 2014a</xref>). The protein sequences were optimized for <italic>Escherichia coli</italic> expression using a multifactor algorithm (<ext-link ext-link-type="uri" xlink:href="https://www.genscript.com/tools/gensmart-codon-optimization">https://www.genscript.com/tools/gensmart-codon-optimization</ext-link>), then synthesized by GenScript. After inserting each gene in pET24a (&#x2b;) via NdeI and NotI restriction enzyme sites, the final plasmids were transformed into NRC <italic>E. coli</italic> BL21-T7 strain (<italic>rhaB lacZ</italic>::P<italic>tac</italic>-T7 RNAP). For each clone, a 2.8-L Fernbach baffled flask containing 500&#xa0;mL Animal-Product Free (APF) LB Miller (Athena Enzyme Systems Cat. 0133) plus 50&#xa0;&#x3bc;g/mL kanamycin was inoculated with an overnight preculture to get an initial OD<sub>600nm</sub> of 0.1. The flasks were incubated at 37&#xb0;C, 200&#x2013;250&#xa0;rpm until an OD<sub>600nm</sub> between 0.8 and 1.0 were reached. To induce protein expression 1&#xa0;mM isopropyl &#x3b2;-d-1-thiogalactopyranoside (IPTG) was added and the culture incubated for another 4&#xa0;h at 37&#xb0;C, 200&#x2013;250&#xa0;rpm. The cultures were harvested, and the cell pellets stored at &#x2212;80&#xb0;C.</p>
<p>Before purification, a cell pellet was resuspended in Lysis buffer 50&#xa0;mM NaPO<sub>4</sub>, 300&#xa0;mM NaCl, 10&#xa0;mM imidazole, pH 7.4 with cOmplete protease inhibitors EDTA-free (Millipore Sigma Cat. 11836170001) and lysed by two passages on a French Pressure Cell Disruptor. Finally, the cell lysate was clarified by centrifugation at 10,000 x <italic>g</italic>, 4&#xb0;C, for 15&#xa0;min and filtration on 0.45&#xa0;&#xb5;m filter. A fraction of the clarified lysate (15&#xa0;mL) was applied on a 3&#xa0;mL HisPur Cobalt Spin Column (Thermo Fisher Cat. 89969) and the column was washed with 20&#xa0;mM NaPO<sub>4</sub>, pH 7.5, 500&#xa0;mM NaCl, 0.3&#xa0;mM TCEP, 15&#xa0;mM imidazole. Elution was done with 20&#xa0;mM NaPO<sub>4</sub>, pH 7.5, 500&#xa0;mM NaCl, 0.3&#xa0;mM TCEP, 100&#xa0;mM imidazole and pooled after visualization on SDS-PAGE. For some of the proteins, the purification was repeated to increase purity. Buffer exchange for DPBS (Thermo Fisher Cat. 14190144) was done with PD-10 desalting columns (Cytiva Cat. 17085101) and final concentration measured by Qubit Protein Assay (Thermo Fisher Cat. Q33211).</p>
<p>The design of BCL-W was based on (<xref ref-type="bibr" rid="B31">Schilling et al., 2014a</xref>) with an N-terminal Avi-tag followed by a bacteriophage lambda protein D fusion tag to improve protein solubility (<xref ref-type="bibr" rid="B11">Forrer and Jaussi, 1998</xref>) (see <xref ref-type="sec" rid="s10">Supplementary Data</xref>). A 6xHis tag was added to the C-terminus of BCL-W for purification. Gene optimization, synthesis and cloning in pET24a (&#x2b;) vector was done as described above for the DARPins. To allow <italic>in vitro</italic> biotinylation, the NRC <italic>E. coli</italic> BL21-T7 strain (<italic>rhaB lacZ</italic>::P<italic>tac</italic>-T7 RNAP) was first transformed with pBirAcm (Avidity), a plasmid expressing biotin ligase under <italic>tac</italic> promoter (IPTG inducible). After growing a chloramphenicol resistant colony in APF LP Miller medium containing 10&#xa0;&#x3bc;g/mL chloramphenicol, electrocompetent cells were prepared using standard procedures. The plasmid pET24a (&#x2b;)-BCL-W was then transformed in BL21-T7/pBirAcm strain and selected on APF LB Miller agar containing 50&#xa0;&#x3bc;g/mL kanamycin and 10&#xa0;&#x3bc;g/mL chloramphenicol.</p>
<p>Expression of BCL-W, cell lysis and clarification were done as described for the DARPins with some exceptions. Both antibiotics, kanamycin and chloramphenicol, were used, and biotin was added to a final concentration of 5&#xa0;mM during the culture (25&#xa0;mL). The cells were lysed in a buffer containing 50&#xa0;mM NaPO<sub>4</sub>, 300&#xa0;mM NaCl, 10&#xa0;mM imidazole, pH 8.0 (plus cOmplete EDTA-free protease inhibitors). The clarified lysate (2.5&#xa0;mL) was applied on a 0.2&#xa0;mL HisPur Cobalt Spin Column (Thermo Fisher Cat. 90090) and the column was washed with 20&#xa0;mM NaPO<sub>4</sub>, pH 7.5, 500&#xa0;mM NaCl, 0.3&#xa0;mM TCEP, 20&#xa0;mM imidazole. Elution was done with 20&#xa0;mM NaPO<sub>4</sub>, pH 7.5, 500&#xa0;mM NaCl, 0.3&#xa0;mM TCEP, 300&#xa0;mM imidazole and pooled after visualization on SDS-PAGE. Buffer exchange for DPBS (Thermo Fisher Cat. 14190144) was done with G-25 MiniTrap desalting columns (Cytiva Cat. 28918007) and final concentration measured by Qubit Protein Assay (Thermo Fisher Cat. Q33211). Purity levels are given in <xref ref-type="sec" rid="s10">Supplementary Table S1</xref> and SDS-PAGE gels are provided as <xref ref-type="sec" rid="s10">Supplementary Data</xref>.</p>
</sec>
<sec id="s2-2-2">
<title>2.2.2 Binding affinity measurements</title>
<p>Surface plasmon resonance was used to screen the top 18 DARPin designs for binding to the biotinylated BCL-W using a Biacore T200 instrument (Cytiva Inc., Marlborough MA) at 25&#xb0;C and with PBST running buffer (Teknova, Hollister CA) containing 0.05% Tween 20, 3.4&#xa0;mM EDTA and an additional 350&#xa0;mM NaCl. The strategy employed was to capture the biotinylated BCL-W onto the SPR surface with a CAP sensor chip (Cytiva Inc.) and flow a three-point concentration series of the DARPin scaffold using a 10-fold dilution series from 1&#xa0;&#x3bc;M to cover a wide concentration range. From the resulting sensorgrams, the affinity constant of binding candidates can be determined. A CAP immobilization chip was prepared following the manufacturer&#x2019;s instructions. Each injection cycle consisted first of a 120-s injection at 5&#xa0;&#x3bc;L/min of a 5-fold dilution of CAP reagent to indirectly immobilize streptavidin over flow-cells 1 and 2. This was followed by a 240-s capture of 5&#xa0;&#x3bc;g/mL biotinylated BCL-W at 5&#xa0;&#x3bc;L/min over flow cell 2 only to form the 60&#x2013;62 RU BCL-W surface, and finally a three-point concentration injection of the DARPin scaffold or running buffer only using single-cycle kinetics was performed at 50&#xa0;&#x3bc;L/min for 90&#xa0;s with a 300-s dissociation phase. At the end of the dissociation phase, any BCL-W/DARPin complex was stripped from the SPR surface using a 60-s injection of 6&#xa0;M GuCl/0.25&#xa0;M NaOH taken from the CAP sensor chip reagent kit. The sensorgrams were double referenced and analyzed using the Biacore BiaEval software. Affinities of the DARPin scaffolds for BCL-W were determined using the steady state model, or the 1:1 binding model when kinetic rate constants could be evaluated.</p>
</sec>
<sec id="s2-2-3">
<title>2.2.3 Folding stability measurements</title>
<p>Differential scanning calorimetry (DSC) was used to determine the thermal transition midpoints (T<sub>m</sub>) as previously performed (<xref ref-type="bibr" rid="B34">Schrag et al., 2019</xref>). DSC was carried out in a VP-Capillary DSC system instrument (Malvern Instruments Ltd., Malvern, United Kingdom). Samples were diluted in DPBS buffer to a final concentration of 0.4&#xa0;mg/mL. DPBS blank and sample scans were carried out by increasing the temperature from 20&#xb0;C to 100&#xb0;C at a rate of 60&#xb0;C/h, with feedback mode/gain set at &#x201c;low&#x201d;, filtering period of 8&#xa0;s, pre-scan time of 3&#xa0;min, and under 70 psi of nitrogen pressure. All data were analyzed with Origin 7.0 software (OriginLab Corporation, Northampton, MA). Thermograms were corrected by subtraction of corresponding DPBS blank scans and normalized to the protein molar concentration. The T<sub>m</sub> values were determined using automated data processing with the rectangular peak finder algorithm for T<sub>m</sub>. Melting temperatures are listed in <xref ref-type="sec" rid="s10">Supplementary Table S1</xref> and DSC thermograms are provided as <xref ref-type="sec" rid="s10">Supplementary Data</xref>.</p>
</sec>
</sec>
</sec>
<sec sec-type="results" id="s3">
<title>3 Results</title>
<sec id="s3-1">
<title>3.1 Sequence-based and structure-based computational design</title>
<sec id="s3-1-1">
<title>3.1.1 Overall design process</title>
<p>The flowchart in <xref ref-type="fig" rid="F1">Figure 1</xref> presents the overall computational design process devised and implemented for this rigid-docking based proof-of-concept engineering study based on the DARPin scaffold. It includes 6 steps: 1) definition of a single DARPin common framework sequence; 2) expansion of the common framework sequence into a DARPin sequence library with variable positions; 3) selection of a small DARPin sub-library consisting of diverse sequences; 4) grafting of the sequence sub-library onto DARPin structural templates; 5) docking of DARPin sub-library to target protein structures, the core component of the process; and 6) ranking docked DARPin variants for experimental testing. The first three steps operate in the sequence space, whereas the last three in the 3D structure space. All the steps are described in detail in the sub-sections of the Methods section. The following sub-sections focus more in-depth on results obtained in steps 3), 4), 5) and 6) of the process.</p>
</sec>
<sec id="s3-1-2">
<title>3.1.2 Selecting diverse DARPin sub-library sequences</title>
<p>Expanding a common framework sequence by varying 18 positions lining the concave face of the DARPin fold (<xref ref-type="fig" rid="F2">Figure 2</xref>) resulted in 10<sup>23</sup> theoretical library size. Millions of iterations were run to select a diverse sub-library fulfilling several design criteria (see Methods sub-<xref ref-type="sec" rid="s2-1-3">Section 2.1.3</xref>). The resulting diverse sub-library comprised a total of 2,213 designs of which 1,429 were produced by mutations and 784 by permutations (<xref ref-type="table" rid="T1">Table 1</xref>). The closest designs in sequence are 9 amino-acid substitutions away from any of the two positive controls (<xref ref-type="sec" rid="s10">Supplementary Figure S1</xref>), or 6 groups away when grouping amino acids by homology (see Methods section). The mutation-based designs have an even proportion of amino-acid groups at the variable positions (<xref ref-type="sec" rid="s10">Supplementary Figure S2</xref>). In contrast, permutation-based designs have unevenly distributed amino-acid groups and lack Ala, His and Ser as inherited from the starting positive-control sequences (<xref ref-type="sec" rid="s10">Supplementary Figure S2</xref>). In terms of net charge, mutation-based designs span a wide range from &#x2212;16 to &#x2b;3 with a mean net charge of &#x2212;6.8, whereas permutation-based designs inherit the net charges of their respective parental positive control (<xref ref-type="sec" rid="s10">Supplementary Figure S3</xref>).</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Library design statistics.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Starting DARPin</th>
<th align="center">PDB ID</th>
<th align="center">Variable positions<xref ref-type="table-fn" rid="Tfn1">
<sup>a</sup>
</xref>
</th>
<th align="center">Set<xref ref-type="table-fn" rid="Tfn2">
<sup>b</sup>
</xref>
</th>
<th align="center">N<sub>seq</sub>
<xref ref-type="table-fn" rid="Tfn3">
<sup>c</sup>
</xref>
</th>
<th align="center">d<sub>seq</sub>
<xref ref-type="table-fn" rid="Tfn4">
<sup>d</sup>
</xref>
</th>
<th align="center">d<sub>chemseq</sub>
<xref ref-type="table-fn" rid="Tfn5">
<sup>e</sup>
</xref>
</th>
<th align="center">Q<sub>net</sub>
<xref ref-type="table-fn" rid="Tfn6">
<sup>f</sup>
</xref>
</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">Common framework</td>
<td align="center">4drx [F]</td>
<td align="center">ASLTYIMSLITWDIMKFK</td>
<td align="center">M</td>
<td align="center">1,429</td>
<td align="center">9</td>
<td align="center">7</td>
<td align="center">&#x2212;6.8</td>
</tr>
<tr>
<td align="center">Known binder</td>
<td align="center">4k5a [B]</td>
<td align="center">KYDMNFMRDNFWKQQKFK</td>
<td align="center">P</td>
<td align="center">284</td>
<td align="center">12</td>
<td align="center">7</td>
<td align="center">&#x2212;4.0</td>
</tr>
<tr>
<td align="center">Known binder</td>
<td align="center">4k5b [B]</td>
<td align="center">RFWMEDLTMKIVYWEKFK</td>
<td align="center">P</td>
<td align="center">500</td>
<td align="center">9</td>
<td align="center">6</td>
<td align="center">&#x2212;6.0</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn id="Tfn1">
<label>
<sup>a</sup>
</label>
<p>Position IDs, in the same order: 45, 46, 48, 56, 57, 78, 79, 81, 89, 90, 111, 112, 114, 122, 123, 144, 145 and 147.</p>
</fn>
<fn id="Tfn2">
<label>
<sup>b</sup>
</label>
<p>M: mutation; P: permutation.</p>
</fn>
<fn id="Tfn3">
<label>
<sup>c</sup>
</label>
<p>Number of sequences.</p>
</fn>
<fn id="Tfn4">
<label>
<sup>d</sup>
</label>
<p>Closest distance from a design to a known binder interface at 18 variable positions, expressed as number of substitutions.</p>
</fn>
<fn id="Tfn5">
<label>
<sup>e</sup>
</label>
<p>Closest distance from a design to a known binder interface at 18 variable positions, expressed as number of homology group changes.</p>
</fn>
<fn id="Tfn6">
<label>
<sup>f</sup>
</label>
<p>Mean net charge of designs within the set.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>In order to generate a sub-library that samples homogeneously the immense theoretical sequence space, designs were imposed to be orthogonal to each other. Clustering based on amino-acid properties indicated that most sequence space regions were covered by both mutation-based and permutation-based types of sequences, with a few areas only covered by the mutation-based set (<xref ref-type="fig" rid="F3">Figure 3</xref>). While proximity in sequence might be perceivable between some of the designs and the two positive controls (<xref ref-type="fig" rid="F3">Figure 3A</xref>), overall, the designed sequences were diverse and nearly equidistant from each other (<xref ref-type="fig" rid="F3">Figure 3B</xref>).</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Diversity of the DARPin sequence sub-library. Unrooted phylogenetic tree from hierarchical clustering of sequences by the chemical properties of amino acids using defined amino-acid homology groups (see Methods section). Sequences marked in black are from the mutation-based set and in blue from the permutation-based set. The two positive-control sequences are shown in red. Only a 5% random sample of the sub-library consisting of 134 sequences is plotted. The top-18 consensus designs and positive controls were annotated. <bold>(A)</bold> For visual clarity, the terminal branches were equally trimmed down to a cladogram giving the illusion of sequence proximity (<xref ref-type="bibr" rid="B41">Yu, 2020</xref>). <bold>(B)</bold> The non-trimmed tree that preserves the ordering in <bold>(A)</bold> is shown to illustrate the true divergence in sequence between designs. For reference, the evolutionary distance is shown.</p>
</caption>
<graphic xlink:href="fmolb-10-1253689-g003.tif"/>
</fig>
</sec>
<sec id="s3-1-3">
<title>3.1.3 Grafting sequence sub-library onto DARPin structural templates</title>
<p>Four crystal structures were used as templates in the modeling of the DARPin ligands (see Methods section). The variance in RMSD among these templates has a mean of 0.95&#xa0;&#xc5;. The template 4drx [F] is more distant due in part to an opening of the last repeated motif of the scaffold. The magnitudes of backbone changes between each of these templates and any of the 2 known DARPin binders of BCL-W are larger than between the 2 known binders (0.42&#xa0;&#xc5;). Thus, backbone RMSDs of 0.91, 0.75, 0.79 and 0.79&#xc5; were calculated to the 4k5a [B] known binder, and of 0.96, 0.75, 0.78 and 0.78&#xc5; to the 4k5b [A] known binder, for the template structures 4drx [F], 4j7w [A], 5le6 [A] and 5lw2 [A], respectively. More backbone variations could be observed in the unstructured region of the fourth ankyrin repeat, where the known BCL-W binders had a distinct conformational topology at the tip of this loop region. These variations in the templates relative to known binders were critical for testing the method in real-life application mode in which the bound backbone structure will be unknown <italic>a priori</italic>.</p>
</sec>
<sec id="s3-1-4">
<title>3.1.4 Docking DARPin sub-library structures to target</title>
<p>The entire set of sequence designs in the selected sub-library was grafted onto four template structures, then cross-docked against two target (BCL-W) structures, leading to 8 docking runs per DARPin sequence. The two backbone structures used for the target (4k5a [A] and 4k5b [C]) were relatively close from each other, with an RMSD of 0.77&#xa0;&#xc5;. They also engaged their respective known DARPin binders (4k5a [B] and 4k5b [A]) via a well-preserved binding interface with backbone atoms deviating by an RMSD of 0.60&#xa0;&#xc5;. Hence, in this study, the docked poses for novel DARPins were required to bind around the same epitope that is targeted by these two known DARPin binders of BCL-W. In more technical terms, the predicted poses of designed DARPins were required to have an overlap of at least acceptable quality (according to CAPRI classification (<xref ref-type="bibr" rid="B21">Lensink et al., 2017</xref>) to either of these known binders. This was met by 1,033 designs (47% of the sub-library), and are referred to as &#x201c;locus designs&#x201d;. (Increasing the stringency and imposing at least a medium quality of pose overlap with the known binders reduced the number of locus designs to 559.) We found no bias towards either of the two target BCL-W structures used for docking, as 811 designs docked to structure 4k5a [A] and 632 designs to structure 4k5b [C]. In terms of the template DARPin structures used for docking, 5lw2 [A] was the least successful template structure with 344 docked designs, followed by 380 designs docked on 5le6 [A], 533 on 4j7w [A] and 579 on 4drx [F]. The net charge distribution of the 1,033 locus designs is slightly different relative the entire docked sub-library of 2,213 designs, as it has sharper peaks at the &#x2212;6 and &#x2212;4 net charges (<xref ref-type="sec" rid="s10">Supplementary Figure S3</xref>).</p>
</sec>
<sec id="s3-1-5">
<title>3.1.5 Ranking DARPin virtual hits</title>
<p>First, locus design DARPins were filtered based on the number of top-1 scored poses, N<sub>pose</sub>, that were docked at the targeted epitope from the 8 docking runs for each DARPin. A total of 293 locus designs (13% of the sub-library) had at least 2 poses docked at the target epitope. These were retained for further ranking and were called &#x201c;consensus designs&#x201d;. The net charge distribution among the consensus designs had even sharper peaks at the net charges &#x2212;6 and &#x2212;4, with the majority of consensus designs at charge &#x2212;6 (<xref ref-type="sec" rid="s10">Supplementary Figure S3</xref>).</p>
<p>For each of selected 293 consensus designs, a consensus score was derived as the arithmetic average over the docking scores of the poses binding to the targeted epitope. These consensus scores were normally distributed and ranged from &#x2212;84.2 to &#x2212;45.3, from strongest to weakest binder (<xref ref-type="fig" rid="F4">Figure 4</xref>). The permutation-based designs were preferentially chosen according to the consensus scores with a median of &#x2212;65 as opposed to a median of &#x2212;61 for the mutation-based ones. In total, 152 (52%) and 71 (24%) designs that docked at the targeted epitope did so with values in N<sub>pose</sub> of 2 and 3, respectively (inset in <xref ref-type="fig" rid="F4">Figure 4</xref>). The lowest consensus score corresponded to a Z-score of &#x2212;3.5. Consensus designs with Z-scores below &#x2212;1.5 were selected for experimental validation, which formed a set consisting of 18 novel DARPins (<xref ref-type="table" rid="T2">Table 2</xref>). An overlay of all consensus poses for the selected designs is shown in <xref ref-type="fig" rid="F5">Figure 5A</xref>.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Distribution of scores from docking-based screening. Distribution of scores obtained from the docking experiments using ProPOSE on the entire set of designs in the library. The scores were obtained from a consensus of multiple predictions binding at the same locus while imposing an acceptable or better quality among the representatives of the cluster. The scores follow a normal distribution with the median marked as dashed lines. The underlying area-under-the-curve of the receiver operating characteristic (AUC-ROC) curve obtained from the separation of the two positive controls from the combined mutation and permutation design sets has a value 0.971. The inset shows the distribution in number of representatives used to calculate the consensus ProPOSE score. The two positive controls have 4 and 7 representatives.</p>
</caption>
<graphic xlink:href="fmolb-10-1253689-g004.tif"/>
</fig>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Top-ranked consensus designs.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Rank</th>
<th align="center">Variable positions<xref ref-type="table-fn" rid="Tfn7">
<sup>a</sup>
</xref>
</th>
<th align="center">Set<xref ref-type="table-fn" rid="Tfn8">
<sup>b</sup>
</xref>
</th>
<th align="center">N<sub>sub</sub>
<xref ref-type="table-fn" rid="Tfn9">
<sup>c</sup>
</xref>
</th>
<th align="center">N<sub>pose</sub>
<xref ref-type="table-fn" rid="Tfn10">
<sup>d</sup>
</xref>
</th>
<th align="center">Q<sub>net</sub>
<xref ref-type="table-fn" rid="Tfn11">
<sup>e</sup>
</xref>
</th>
<th align="center">Score<xref ref-type="table-fn" rid="Tfn12">
<sup>f</sup>
</xref>
</th>
<th align="center">Z-Score<xref ref-type="table-fn" rid="Tfn13">
<sup>g</sup>
</xref>
</th>
<th align="center">K<sub>D</sub> (nM)<xref ref-type="table-fn" rid="Tfn14">
<sup>h</sup>
</xref>
</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">1</td>
<td align="center">RMTKEKFFWEILWYDMVK</td>
<td align="center">P</td>
<td align="center">14</td>
<td align="center">7</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;84.2</td>
<td align="center">&#x2212;3.5</td>
<td align="center">weak</td>
</tr>
<tr>
<td align="center">2</td>
<td align="center">RQIVHRHWFDVIKYWRHL</td>
<td align="center">M</td>
<td align="center">18 (17)</td>
<td align="center">3</td>
<td align="center">&#x2212;1</td>
<td align="center">&#x2212;77.8</td>
<td align="center">&#x2212;2.4</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="center">3</td>
<td align="center">KFWFETMDKMKRYEWVIL</td>
<td align="center">P</td>
<td align="center">14</td>
<td align="center">7</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;77.2</td>
<td align="center">&#x2212;2.3</td>
<td align="center">weak</td>
</tr>
<tr>
<td align="center">4</td>
<td align="center">KFWMEMLTDWIYEVRKKF</td>
<td align="center">P</td>
<td align="center">10</td>
<td align="center">3</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;75.9</td>
<td align="center">&#x2212;2.1</td>
<td align="center">44</td>
</tr>
<tr>
<td align="center">5</td>
<td align="center">KFMREEFWWLIKKTDYMV</td>
<td align="center">P</td>
<td align="center">15</td>
<td align="center">6</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;75.3</td>
<td align="center">&#x2212;2.0</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="center">6</td>
<td align="center">KFWYNDFQMDFQMRNKKK</td>
<td align="center">P</td>
<td align="center">13</td>
<td align="center">4</td>
<td align="center">&#x2212;4</td>
<td align="center">&#x2212;75.2</td>
<td align="center">&#x2212;2.0</td>
<td align="center">150</td>
</tr>
<tr>
<td align="center">7</td>
<td align="center">VWWEEDFKIKMMKFYTLR</td>
<td align="center">P</td>
<td align="center">14</td>
<td align="center">3</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;75.1</td>
<td align="center">&#x2212;2.0</td>
<td align="center">111</td>
</tr>
<tr>
<td align="center">8</td>
<td align="center">KYRKNKFWFNDQFKDQMM</td>
<td align="center">P</td>
<td align="center">14</td>
<td align="center">3</td>
<td align="center">&#x2212;4</td>
<td align="center">&#x2212;74.8</td>
<td align="center">&#x2212;1.9</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="center">9</td>
<td align="center">RKMDQKFKMNDYWNFQFK</td>
<td align="center">P</td>
<td align="center">15</td>
<td align="center">2</td>
<td align="center">&#x2212;4</td>
<td align="center">&#x2212;74.7</td>
<td align="center">&#x2212;1.9</td>
<td align="center">weak</td>
</tr>
<tr>
<td align="center">10</td>
<td align="center">KIMWFKWDYKELMVETFR</td>
<td align="center">P</td>
<td align="center">15</td>
<td align="center">3</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;74.7</td>
<td align="center">&#x2212;1.9</td>
<td align="center">weak</td>
</tr>
<tr>
<td align="center">11</td>
<td align="center">RAVNRTVFVYWAYNFRVV</td>
<td align="center">M</td>
<td align="center">18 (16)</td>
<td align="center">2</td>
<td align="center">&#x2212;4</td>
<td align="center">&#x2212;74.6</td>
<td align="center">&#x2212;1.9</td>
<td align="center">weak</td>
</tr>
<tr>
<td align="center">12</td>
<td align="center">KFWMQRFMQYKDFKDKNN</td>
<td align="center">P</td>
<td align="center">15</td>
<td align="center">2</td>
<td align="center">&#x2212;4</td>
<td align="center">&#x2212;74.6</td>
<td align="center">&#x2212;1.9</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="center">13</td>
<td align="center">KLMEYDFMVWITKFERWK</td>
<td align="center">P</td>
<td align="center">14</td>
<td align="center">7</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;74.6</td>
<td align="center">&#x2212;1.7</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="center">14</td>
<td align="center">KYWYRTTWYHAIWNFYKQ</td>
<td align="center">M</td>
<td align="center">18 (16)</td>
<td align="center">5</td>
<td align="center">&#x2212;3</td>
<td align="center">&#x2212;73.5</td>
<td align="center">&#x2212;1.7</td>
<td align="center">weak</td>
</tr>
<tr>
<td align="center">15</td>
<td align="center">KYFEWVQRVMFKVVLMNR</td>
<td align="center">M</td>
<td align="center">18 (14)</td>
<td align="center">2</td>
<td align="center">&#x2212;4</td>
<td align="center">&#x2212;73.3</td>
<td align="center">&#x2212;1.7</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="center">16</td>
<td align="center">FKMWEMLFWRVIYEDKKT</td>
<td align="center">P</td>
<td align="center">14</td>
<td align="center">5</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;72.8</td>
<td align="center">&#x2212;1.6</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="center">17</td>
<td align="center">KFFRNNKMDYWKKMDFQQ</td>
<td align="center">P</td>
<td align="center">14</td>
<td align="center">3</td>
<td align="center">&#x2212;4</td>
<td align="center">&#x2212;72.8</td>
<td align="center">&#x2212;1.6</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="center">18</td>
<td align="center">KKSQTSYHHQQMLRTHRV</td>
<td align="center">M</td>
<td align="center">18 (17)</td>
<td align="center">5</td>
<td align="center">0</td>
<td align="center">&#x2212;72.7</td>
<td align="center">&#x2212;1.6</td>
<td align="center">n.d.b</td>
</tr>
<tr>
<td align="left"/>
<td align="center">KYDMNFMRDNFWKQQKFK</td>
<td align="center">PC1</td>
<td align="center">0</td>
<td align="center">4</td>
<td align="center">&#x2212;4</td>
<td align="center">&#x2212;75.8</td>
<td align="center">&#x2212;2.1</td>
<td align="center">240</td>
</tr>
<tr>
<td align="left"/>
<td align="center">RFWMEDLTMKIVYWEKFK</td>
<td align="center">PC2</td>
<td align="center">0</td>
<td align="center">7</td>
<td align="center">&#x2212;6</td>
<td align="center">&#x2212;73.5</td>
<td align="center">&#x2212;1.7</td>
<td align="center">0.9</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn id="Tfn7">
<label>
<sup>a</sup>
</label>
<p>Position IDs, in the same order: 45, 46, 48, 56, 57, 78, 79, 81, 89, 90, 111, 112, 114, 122, 123, 144, 145 and 147.</p>
</fn>
<fn id="Tfn8">
<label>
<sup>b</sup>
</label>
<p>P: permutation; M: mutation; PC: positive control.</p>
</fn>
<fn id="Tfn9">
<label>
<sup>c</sup>
</label>
<p>Number of substitutions at 18 variable positions from the corresponding known binder for the P-set designs or from the initial sequence of the common framework-based library for the M-set designs. Number of substitutions from the closest known binder is also shown in parenthesis for the M-set designs.</p>
</fn>
<fn id="Tfn10">
<label>
<sup>d</sup>
</label>
<p>Number of poses predicted to bind at the target epitope.</p>
</fn>
<fn id="Tfn11">
<label>
<sup>e</sup>
</label>
<p>Net charge.</p>
</fn>
<fn id="Tfn12">
<label>
<sup>f</sup>
</label>
<p>Consensus docking score obtained from an arithmetic average of the docked poses at target epitope.</p>
</fn>
<fn id="Tfn13">
<label>
<sup>g</sup>
</label>
<p>Calculated from scores over the set of 293 &#x201c;consensus designs&#x201d; (see Results section).</p>
</fn>
<fn id="Tfn14">
<label>
<sup>h</sup>
</label>
<p>Determined by SPR measurements (see Methods section); weak: K<sub>D</sub> &#x3e; 1&#xa0;&#x3bc;M; n.d.b.: no detected binding.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Non-cognate docking results for the top-ranked poses. <bold>(A)</bold> Overview of all poses at the target epitope for the top-18 consensus designs selected for testing. The novel designs are part of the permutation-based set (blue) and the mutation-based set (black). <bold>(B)</bold> Overview of all poses docked at the target epitope for the two positive controls (red). Comparisons of atomic details between the best-scored docked pose of a positive control and the crystal structure of the corresponding known binder sharing the same residues at the 18 variable positions are shown in panel <bold>(C)</bold> for the positive control PC1 and the known binder (4k5a [A]; in purple), and in panel <bold>(D)</bold> for the positive control PC2 and the known binder (4k5b [C]; in purple). All structure orientations are kept as in <xref ref-type="fig" rid="F2">Figure 2C</xref>.</p>
</caption>
<graphic xlink:href="fmolb-10-1253689-g005.tif"/>
</fig>
<p>The two positive-control interfaces, having the 18 variable positions imported from the two known DARPin binders of BCL-W and grafted onto the common framework sequence, were also docked in the same manner. These positive controls had consensus scores of &#x2212;75.8 and &#x2212;73.5, corresponding to Z-scores of &#x2212;2.1 and &#x2212;1.7, respectively. With the assumption that none of the designs are true binders, the separation of the positive controls from the designs had an AUC of 0.971 (<xref ref-type="fig" rid="F4">Figure 4</xref>). The AUC dropped to 0.922 when best scores were used instead of consensus scores for designs with N<sub>pose</sub> &#x3e; 1. While working with rigid scaffolds, this observation suggested the need for structural ensembles to achieve better enrichments and thus motivated the use of consensus scores over best scores in the sections that follow. These positive controls were ranked within the range of the top-18 novel designs, and they were also subjected to experimental testing. It is important to note that to properly compare the scores for the two parental known binders of BCL-W with those of the mutants, we needed to base it on the modeled structures of the known binders rather than their crystal structures. Using the crystal structures would be a case of cognate backbone docking and perfect match in shape complementarity leading to out-of-range scores (DARPin/BCL-W docking scores of &#x2212;144.5 and &#x2212;123.3 were obtained for 4k5a [B]/4k5a [A] and 4k5b [B]/4k5b [C], respectively). The overlay of all locus docked poses for the two grafted positive control interfaces is shown in <xref ref-type="fig" rid="F5">Figure 5B</xref> to have the same orientation with those of the selected designs (<xref ref-type="fig" rid="F5">Figure 5A</xref>). These poses are further similarly oriented with those of the known binders, as exemplified in <xref ref-type="fig" rid="F5">Figure 5C</xref>. A closer examination reveals that despite an excellent pose recovery for this cross-docking experiment, there are certain noticeable differences in the fine atomic details at the interface, which are likely due mainly to non-cognate backbone coordinates and to a lesser extent to changes of the framework sequence outside the 18 variable positions. Overall, cross-docking of positive controls predicted that they would retain similar binding relative to the corresponding known binders.</p>
<p>As presented in <xref ref-type="table" rid="T2">Table 2</xref>, most of these top consensus designs (13 of 18) were from the permutation set despite its smaller representation in the initial library. Also, 15 out of the 18 consensus designs had a net charge equal to that of a positive control (&#x2212;6 or &#x2212;4), despite the random sequence generation procedure employed. On average, the top-18 designs were 15 mutations away from the positive-control interfaces, with the closest design being 10 mutations away. These top consensus designs had between 2 and 7 top-1 poses bound at the target epitope. Interestingly, a strong bias towards an increased consensus was observed with 14 out of the 18 novel designs (78%) with at least 3 representative poses bound at the targeted epitope. This level has to be contrasted to only 24% of all consensus designs reaching an N<sub>pose</sub> &#x3e; 2. Hence, not only were the top designs predicted to bind stronger to the target, they also did so with a higher number of predicted consensus poses, with an average N<sub>pose</sub> of 4.2. Comparably, the two positive controls had N<sub>pose</sub> values of 4 and 7 (<xref ref-type="fig" rid="F4">Figure 4</xref>).</p>
</sec>
</sec>
<sec id="s3-2">
<title>3.2 Experimental testing of DARPin designs</title>
<p>The 18 top-ranked consensus designs, together with the 2 positive controls and the 2 parental known binder DARPins were produced in bacteria, purified by IMAC and screened for binding to BCL-W by SPR. The purity levels of the DARPins ranged from 45% to 99% with an average of 82% (<xref ref-type="sec" rid="s10">Supplementary Table S1</xref>). While some of these levels could be considered as suboptimal for SPR experiments and might lead to non-specific binding, they were deemed sufficient for a first-pass screening. Tested DARPins were flowed at a fixed concentration over biotinylated target protein immobilized on the sensorchip. An overview of the SPR binding screen is given in <xref ref-type="fig" rid="F6">Figure 6</xref>. Overall, binding in the nM range was detected for 3 designs, the 2 positive controls and the 2 known binders (<xref ref-type="table" rid="T2">Table 2</xref>). Additionally, 6 designs had weak binding in the &#x3bc;M range, with a caveat that some of the binding events detected in these cases could be non-specific. Among the top 10 designs, only 3 had no detected binding, while the 3 stronger binders and 4 of the weak binders were present in this group. All 7 binders in the top-10 group belonged to the permutation (P) set. In the group consisting of the 8 remaining tested designs, ranks 11&#x2013;18, there were only 2 weak binders while the rest of designs had no detected binding. These 2 weak binders were both from the mutation (M) set. Overall, data in <xref ref-type="table" rid="T2">Table 2</xref> indicate a certain level of enrichment in binding that follows the predicted docking scores within the set of 18 tested variants, with the caveat that SPR data is insufficient to confirm the predicted binding modes. We also measured the thermal stabilities of the designed DARPins and obtained very high thermostabilities, with melting temperature (T<sub>m</sub>) values typically in the 80&#x2013;100&#xb0;C range (<xref ref-type="sec" rid="s10">Supplementary Table S1</xref>), comparable with those measured here for the positive controls and known binders, as well as previously for other DARPins (<xref ref-type="bibr" rid="B31">Schilling et al., 2014a</xref>). This stability data provides some level of confidence that the sequence perturbations introduced in the designed variants were able to maintain the folded structure of the archetypical DARPin scaffold.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Surface plasmon resonance screening. SPR binding sensorgrams are shown for the 18 top-ranked designs, the positive controls and the known binders. Ranking of designs is based on the consensus score (see also <xref ref-type="table" rid="T2">Table 2</xref>). Sensorgrams are labeled according to 3 levels of binding affinity as shown in the legend.</p>
</caption>
<graphic xlink:href="fmolb-10-1253689-g006.tif"/>
</fig>
<p>For the two known DARPin binders of BCL-W, 4k5a [B] and 4k5b [B], we obtained dissociation constants, K<sub>D</sub>, of 26&#xa0;nM and 3.5&#xa0;nM, respectively, which are in line with their previously published K<sub>D</sub> data of 10&#xa0;nM and 0.64&#xa0;nM (<xref ref-type="bibr" rid="B31">Schilling et al., 2014a</xref>). The two corresponding positive control DARPins, which import only the 18 variable positions of the common framework scaffold from these known binders, bound with K<sub>D</sub> values of 245&#xa0;nM and 0.9&#xa0;nM. These values represent comparable affinities to their respective parental known binders, although it seems that the framework change from the known binders to the common framework sequence impacted detrimentally the 4k5a [B] interface and beneficially the 4k5b [B] interface.</p>
<p>For the 3 novel DARPin designs exhibiting good binding, we obtained dissociation constants, K<sub>D</sub>, in the 40&#x2013;150&#xa0;nM range, which are well within the range bracketed by the two positive controls (0.9&#x2013;245&#xa0;nM). These were ranked 4, 6 and 7 among the top-18 consensus designs (<xref ref-type="table" rid="T2">Table 2</xref>), with the 4<sup>th</sup> ranked design exhibiting the better K<sub>D</sub> of 44&#xa0;nM, which is similar to the affinity of one of the known binders (26&#xa0;nM). Low binding, with K<sub>D</sub> above 1&#xa0;&#x3bc;M, could also be detected for designs with ranks 1, 3, 9, 10, 11 and 14. Further details about the binders on their amino-acid substitutions, net charges, sets and substitutions are listed in <xref ref-type="table" rid="T2">Table 2</xref>.</p>
<p>A retrospective analysis of ranking by best scores instead of consensus scores versus experiment indicated that this approach could also be suitable (<xref ref-type="sec" rid="s10">Supplementary Table S2</xref>). By this ranking of the 18 tested designs, the top 4 gave binding signals and among them the 2<sup>nd</sup> and 4<sup>th</sup> ranked are the best designs with K<sub>D</sub> values of 44&#xa0;nM and 111&#xa0;nM. Also, best scoring was able to correctly rank the two positive controls among themselves, i.e., the stronger binder has a more negative score. However, best scoring always performed slightly worse than consensus ranking for the discrimination of binders against non-binders (<xref ref-type="sec" rid="s10">Supplementary Figure S4</xref>).</p>
<p>While our computational strategy forced the designs to bind at a specified locus, our geometric criteria were loose enough to allow for some structural variability around the targeted epitope, that could lead to substantially different structural determinants required for binding. Despite the weak statistics due to the relatively low number of experimentally-validated designs, a close inspection of important structural determinants revealed that the non-binders bury more surface area on average than the validated strong binders (<xref ref-type="sec" rid="s10">Supplementary Figure S5</xref>). Notably, a larger fraction in non-polar surface area on the BCL-W interface is predicted to be lost by the non-binders relative to binders (<xref ref-type="sec" rid="s10">Supplementary Figure S5</xref>). This is an interesting finding to explore in future screening campaigns as docking algorithms are normally calibrated to attribute larger scores to burial of larger interfaces and would indirectly favor or enrich those designs achieving increased surface burial. For this set of binders, hydrophobic residues tend to be preferentially enriched only in the internal DARPin repeat 1 (<xref ref-type="sec" rid="s10">Supplementary Figure S5</xref>).</p>
</sec>
</sec>
<sec sec-type="discussion" id="s4">
<title>4 Discussion</title>
<p>In this proof-of-concept study, we aimed at exploring if rigid-backbone docking can lead to meaningful biologics discovery. A first objective was to test, in a real-life scenario, the utility of our exhaustive protein-protein docking tool ProPOSE that incorporates side-chain flexibility (<xref ref-type="bibr" rid="B13">Hogues et al., 2018</xref>). ProPOSE performed very well in cognate-backbone docking, but returned a lower performance in unbound-backbone docking, thus hampering <italic>de novo</italic> antibody discovery efforts, mainly due to the hypervariable nature of the CDR-H3 loop. While work addressing the challenging problem of backbone sampling and scoring is highly relevant and remains to be pursued, here we explored the practical utility of ProPOSE in its current state by employing a more rigid scaffold, DARPin, which has already been used as an alternative scaffold in biologics discovery (<xref ref-type="bibr" rid="B4">Binz et al., 2003</xref>; <xref ref-type="bibr" rid="B24">Pluckthun, 2015</xref>). The overarching assumption is that ProPOSE can tolerate some minor level of backbone movements at the binding interface, but the extent of tolerated backbone movements has not been established yet.</p>
<p>From the technological perspective of rigid docking with unbound backbone conformation, employing four experimentally determined backbone conformations, each slightly different from bound backbone conformations, provided a test of the impact of backbone flexibility on biologics design. An initial measure of success was gleaned from so-called positive controls, in which 18 interfacial residues of known DARPin binders to a given target (BCL-W in this study) were transferred to a common DARPin framework sequence, assigned unbound backbone conformations, and cross-docked to the target. The predicted binding modes of these positive controls were similar to those of known binders, but docking scores were reduced almost in half relative to those obtained for the known binders in their bound backbone conformations. Yet, experimental testing of these positive controls showed retained binding affinities at comparable levels relative to the known binders, despite reduced scores. This established a new range of binding scores at a reduced magnitude which was adapted for cross-docking but remained predictive of true-positive binders. Consequently, novel DARPin designs cross-docked at that same target epitope were top-ranked and had scores within the re-established score level suitable for cross-docking. Upon their experimental testing, seven out of top-10 ranked designs demonstrated at least some level of binding to the target, with 3 of them exhibiting binding strengths similar to those of the positive controls as well as the previously known binders.</p>
<p>Despite this initial relative success, rigid-backbone docking remains challenging even for scaffolds with fairly rigid protein backbone like DARPins. Several limitations of this approach and directions for possible improvements are noted below.</p>
<p>First, most novel binders belonged to the random permutation (P) set, which confines the library space with respect to certain global properties, for example, the net charge. These results thus point to the benefits of landing into the &#x201c;right&#x201d; regions of the library space after randomization at variable positions. While it was certainly harder for members from the random mutation (M) set to reach the top of the hit list, the finding of two weak binders belonging the M-set is extremely encouraging. In principle, real-life applications utilize mainly M-libraries. One way in this direction could be to enlarge the size of the docked diversified sub-library (only &#x223c;2,000 in this study). This could be feasible with access to large computing resources given the not overly prohibitive computational task involved in running ProPOSE. An alternative approach could be a focused expansion into P-subsets around initial M-set hits from a relatively sparse sub-library. This approach could set a preferred range for net charge, for example, and it would be especially beneficial as the number of randomized interfacial positions increases. Furthermore, the efficiency of the M-libraries at finding better hits could most likely be improved by applying structure-guided filters to search in more relevant regions of the sequence space. For instance, designs could be filtered based on their complementarity in charge or by their exposure of polar or non-polar surfaces at variable positions based on structural information of the selected binding epitope being targeted.</p>
<p>Secondly, while initial hits are often weak binders which are difficult to characterize, they should not be immediately discarded but rather treated as seeds for further optimization by affinity maturation, which can be done either experimentally (e.g., display methods) or computationally (e.g., ADAPT platform). This aspect has significant practical importance, given that by random sampling of the immense library space it is highly unlikely to obtain a very strong binder.</p>
<p>Thirdly, the unbound backbone conformations selected for cross-docking were from experimentally determined crystal structures. This is similar to the multiple protein structure approach used in small-molecule docking and virtual screening (<xref ref-type="bibr" rid="B35">Sheridan et al., 2008</xref>). Because the DARPin scaffold is not completely rigid and scoring functions used in docking are sensitive to atomic positions, including more than one backbone as templates in the cross-docking approach was felt to be beneficial. Carefully derived simulated structures obtained, for example, via backrub motions, molecular dynamics or Monte-Carlo simulations can be used as alternatives sources to experimentally-determined backbone conformations. The multiple template approach used here for docking was also extended to the stage of hit ranking, via consensus scoring. This seemed to provide a reasonable enrichment, although retrospectively we also found that the best-score approach might provide a similarly good, if not better ranking, among the small set of hits ranked by consensus scoring.</p>
<p>Despite some approximations in the underlying methodology adopted here, it is encouraging that cross-docking could identify binding sequences that differ substantially from known binders out of thousands of potential candidates. This relative success may be attributed to the foundational work underlying the methods used here to address the two intimately-related challenges of docking and scoring in computational drug discovery (<xref ref-type="bibr" rid="B33">Schneider et al., 2022</xref>). On one hand, for binding mode prediction, ProPOSE was used given its high accuracy in rigid-backbone docking when the bound-backbone conformation is provided. On the other hand, for ranking among different docked variants, ProPOSE employed a scoring function drawn from the solvated interaction energy (SIE) exhibiting high transferability from small-molecule to protein ligands (<xref ref-type="bibr" rid="B25">Purisima et al., 2023</xref>).</p>
<p>The data presented here support the notion that <italic>de novo</italic> biologics discovery <italic>via</italic> computational methods is a tractable problem that could complement the more traditional and matured wet-lab methods of library display screening and animal immunization. One main added benefit of the structure-based approach is directing the binding response towards desired target locations, e.g., functionally relevant, in a controlled manner. Further advances in several areas such as backbone sampling and depth of theoretical library screening, will be required for maturing <italic>de novo</italic> biologics discovery for routine applications in the not-so-distant future.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/<xref ref-type="sec" rid="s10">Supplementary Material</xref>, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec id="s6">
<title>Author contributions</title>
<p>FG: Conceptualization, Data curation, Formal Analysis, Investigation, Methodology, Software, Visualization, Writing&#x2013;original draft, Writing&#x2013;review and editing. JB: Formal Analysis, Investigation, Methodology, Visualization, Writing&#x2013;original draft, Writing&#x2013;review and editing. YM: Formal Analysis, Investigation, Methodology, Writing&#x2013;review and editing. AD: Formal Analysis, Investigation, Methodology, Writing&#x2013;review and editing. HH. Conceptualization, Writing&#x2013;review and editing. CC. Conceptualization, Writing&#x2013;review and editing. EP: Conceptualization, Resources, Supervision, Writing&#x2013;review and editing. MA: Formal Analysis, Methodology, Resources, Supervision, Writing&#x2013;original draft, Writing&#x2013;review and editing. TS: Conceptualization, Formal Analysis, Project administration, Resources, Supervision, Visualization, Writing&#x2013;original draft, Writing&#x2013;review and editing.</p>
</sec>
<sec id="s7">
<title>Funding</title>
<p>Some of the docking simulations were run on high-performance computing clusters from Digital Research Alliance of Canada under project number 4191 (TS).</p>
</sec>
<ack>
<p>We thank Emma Smith for performing DSC experiments.</p>
</ack>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s9">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="s10">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fmolb.2023.1253689/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fmolb.2023.1253689/full&#x23;supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="DataSheet1.PDF" id="SM1" mimetype="application/PDF" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Abanades</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Wong</surname>
<given-names>W. K.</given-names>
</name>
<name>
<surname>Boyles</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Georges</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Bujotzek</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Deane</surname>
<given-names>C. M.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>ImmuneBuilder: deep-Learning models for predicting the structures of immune proteins</article-title>, <comment>2011.2004. bioRxiv, 514231.</comment> <pub-id pub-id-type="doi">10.1101/2022.11.04.514231</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Adolf-Bryfogle</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Kalyuzhniy</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Kubitz</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Weitzner</surname>
<given-names>B. D.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Adachi</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>RosettaAntibodyDesign (RAbD): a general framework for computational antibody design</article-title>. <source>PLoS Comput. Biol.</source> <volume>14</volume>, <fpage>e1006112</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pcbi.1006112</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Berman</surname>
<given-names>H. M.</given-names>
</name>
<name>
<surname>Westbrook</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Feng</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Gilliland</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Bhat</surname>
<given-names>T. N.</given-names>
</name>
<name>
<surname>Weissig</surname>
<given-names>H.</given-names>
</name>
<etal/>
</person-group> (<year>2000</year>). <article-title>The protein Data Bank</article-title>. <source>Nucl. Acids Res.</source> <volume>28</volume>, <fpage>235</fpage>&#x2013;<lpage>242</lpage>. <pub-id pub-id-type="doi">10.1093/nar/28.1.235</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Binz</surname>
<given-names>H. K.</given-names>
</name>
<name>
<surname>Stumpp</surname>
<given-names>M. T.</given-names>
</name>
<name>
<surname>Forrer</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Amstutz</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Pluckthun</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2003</year>). <article-title>Designing repeat proteins: well-expressed, soluble and stable proteins from combinatorial libraries of consensus ankyrin repeat proteins</article-title>. <source>J. Mol. Biol.</source> <volume>332</volume>, <fpage>489</fpage>&#x2013;<lpage>503</lpage>. <pub-id pub-id-type="doi">10.1016/s0022-2836(03)00896-9</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chowdhury</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Allan</surname>
<given-names>M. F.</given-names>
</name>
<name>
<surname>Maranas</surname>
<given-names>C. D.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>OptMAVEn-2.0: de novo design of variable antibody regions against targeted antigen epitopes</article-title>. <source>Antibodies (Basel)</source> <volume>7</volume>, <fpage>23</fpage>. <pub-id pub-id-type="doi">10.3390/antib7030023</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cohen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Halfon</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Schneidman-Duhovny</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>NanoNet: rapid and accurate end-to-end nanobody modeling by deep learning</article-title>. <source>Front. Immunol.</source> <volume>13</volume>, <fpage>958584</fpage>. <pub-id pub-id-type="doi">10.3389/fimmu.2022.958584</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cornell</surname>
<given-names>W. D.</given-names>
</name>
<name>
<surname>Cieplak</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Bayly</surname>
<given-names>C. I.</given-names>
</name>
<name>
<surname>Gould</surname>
<given-names>I. R.</given-names>
</name>
<name>
<surname>Merz</surname>
<given-names>K. M.</given-names>
</name>
<name>
<surname>Ferguson</surname>
<given-names>D. M.</given-names>
</name>
<etal/>
</person-group> (<year>1995</year>). <article-title>A second generation force field for the simulation of proteins, nucleic acids, and organic molecules</article-title>. <source>J. Am. Chem. Soc.</source> <volume>117</volume>, <fpage>5179</fpage>&#x2013;<lpage>5197</lpage>. <pub-id pub-id-type="doi">10.1021/ja00124a002</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>DeFrancesco</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Drug pipeline 1Q19</article-title>. <source>Nat. Biotechnol.</source> <volume>37</volume>, <fpage>579</fpage>&#x2013;<lpage>580</lpage>. <pub-id pub-id-type="doi">10.1038/s41587-019-0146-7</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Evans</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>O&#x2019;Neill</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Pritzel</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Antropova</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Senior</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Green</surname>
<given-names>T.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Protein complex prediction with AlphaFold-Multimer</article-title>, <comment>2010.2004.463034. bioRxiv</comment>. <pub-id pub-id-type="doi">10.1101/2021.10.04.463034</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fischman</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ofran</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Computational design of antibodies</article-title>. <source>Curr. Opin. Struct. Biol.</source> <volume>51</volume>, <fpage>156</fpage>&#x2013;<lpage>162</lpage>. <pub-id pub-id-type="doi">10.1016/j.sbi.2018.04.007</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Forrer</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Jaussi</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>1998</year>). <article-title>High-level expression of soluble heterologous proteins in the cytoplasm of <italic>Escherichia coli</italic> by fusion to the bacteriophage lambda head protein D</article-title>. <source>Gene</source> <volume>224</volume>, <fpage>45</fpage>&#x2013;<lpage>52</lpage>. <pub-id pub-id-type="doi">10.1016/s0378-1119(98)00538-1</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gao</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Nakajima An</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Parks</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Skolnick</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>AF2Complex predicts direct physical interactions in multimeric proteins with deep learning</article-title>. <source>Nat. Commun.</source> <volume>13</volume>, <fpage>1744</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-022-29394-2</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hogues</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Gaudreault</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Corbeil</surname>
<given-names>C. R.</given-names>
</name>
<name>
<surname>Deprez</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Sulea</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Purisima</surname>
<given-names>E. O.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>ProPOSE: direct exhaustive protein-protein docking with side chain flexibility</article-title>. <source>J. Chem. Theory Comput.</source> <volume>14</volume>, <fpage>4938</fpage>&#x2013;<lpage>4947</lpage>. <pub-id pub-id-type="doi">10.1021/acs.jctc.8b00225</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hogues</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Sulea</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Purisima</surname>
<given-names>E. O.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Exhaustive docking and solvated interaction energy scoring: lessons learned from the SAMPL4 challenge</article-title>. <source>J. Comput. Aided Mol. Des.</source> <volume>28</volume>, <fpage>417</fpage>&#x2013;<lpage>427</lpage>. <pub-id pub-id-type="doi">10.1007/s10822-014-9715-5</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hornak</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Abel</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Okur</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Strockbine</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Roitberg</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Simmerling</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>Comparison of multiple Amber force fields and development of improved protein backbone parameters</article-title>. <source>Proteins</source> <volume>65</volume>, <fpage>712</fpage>&#x2013;<lpage>725</lpage>. <pub-id pub-id-type="doi">10.1002/prot.21123</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jumper</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Evans</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Pritzel</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Green</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Figurnov</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Ronneberger</surname>
<given-names>O.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Highly accurate protein structure prediction with AlphaFold</article-title>. <source>Nature</source> <volume>596</volume>, <fpage>583</fpage>&#x2013;<lpage>589</lpage>. <pub-id pub-id-type="doi">10.1038/s41586-021-03819-2</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kaplon</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Crescioli</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Chenoweth</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Visweswaraiah</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Reichert</surname>
<given-names>J. M.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Antibodies to watch in 2023</article-title>. <source>MAbs</source> <volume>15</volume>, <fpage>2153410</fpage>. <pub-id pub-id-type="doi">10.1080/19420862.2022.2153410</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kramer</surname>
<given-names>M. A.</given-names>
</name>
<name>
<surname>Wetzel</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Pluckthun</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Mittl</surname>
<given-names>P. R.</given-names>
</name>
<name>
<surname>Grutter</surname>
<given-names>M. G.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Structural determinants for improved stability of designed ankyrin repeat proteins with a redesigned C-capping module</article-title>. <source>J. Mol. Biol.</source> <volume>404</volume>, <fpage>381</fpage>&#x2013;<lpage>391</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmb.2010.09.023</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Krivov</surname>
<given-names>G. G.</given-names>
</name>
<name>
<surname>Shapovalov</surname>
<given-names>M. V.</given-names>
</name>
<name>
<surname>Dunbrack</surname>
<given-names>R. L.</given-names>
<suffix>Jr.</suffix>
</name>
</person-group> (<year>2009</year>). <article-title>Improved prediction of protein side-chain conformations with SCWRL4</article-title>. <source>Proteins</source> <volume>77</volume>, <fpage>778</fpage>&#x2013;<lpage>795</lpage>. <pub-id pub-id-type="doi">10.1002/prot.22488</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Larkin</surname>
<given-names>M. A.</given-names>
</name>
<name>
<surname>Blackshields</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Brown</surname>
<given-names>N. P.</given-names>
</name>
<name>
<surname>Chenna</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>McGettigan</surname>
<given-names>P. A.</given-names>
</name>
<name>
<surname>McWilliam</surname>
<given-names>H.</given-names>
</name>
<etal/>
</person-group> (<year>2007</year>). <article-title>Clustal W and clustal X version 2.0</article-title>. <source>Bioinformatics</source> <volume>23</volume>, <fpage>2947</fpage>&#x2013;<lpage>2948</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btm404</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lensink</surname>
<given-names>M. F.</given-names>
</name>
<name>
<surname>Velankar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Wodak</surname>
<given-names>S. J.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Modeling protein-protein and protein-peptide complexes: CAPRI 6th edition</article-title>. <source>Proteins</source> <volume>85</volume>, <fpage>359</fpage>&#x2013;<lpage>377</lpage>. <pub-id pub-id-type="doi">10.1002/prot.25215</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lu</surname>
<given-names>R. M.</given-names>
</name>
<name>
<surname>Hwang</surname>
<given-names>Y. C.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>I. J.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>C. C.</given-names>
</name>
<name>
<surname>Tsai</surname>
<given-names>H. Z.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>H. J.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Development of therapeutic antibodies for the treatment of diseases</article-title>. <source>J. Biomed. Sci.</source> <volume>27</volume>, <fpage>1</fpage>. <pub-id pub-id-type="doi">10.1186/s12929-019-0592-z</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pecqueur</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Duellberg</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Dreier</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Jiang</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Pluckthun</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2012</year>). <article-title>A designed ankyrin repeat protein selected to bind to tubulin caps the microtubule plus end</article-title>. <source>Proc. Natl. Acad. Sci. U. S. A.</source> <volume>109</volume>, <fpage>12011</fpage>&#x2013;<lpage>12016</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.1204129109</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pluckthun</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Designed ankyrin repeat proteins (DARPins): binding proteins for research, diagnostics, and therapy</article-title>. <source>Annu. Rev. Pharmacol. Toxicol.</source> <volume>55</volume>, <fpage>489</fpage>&#x2013;<lpage>511</lpage>. <pub-id pub-id-type="doi">10.1146/annurev-pharmtox-010611-134654</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Purisima</surname>
<given-names>E. O.</given-names>
</name>
<name>
<surname>Corbeil</surname>
<given-names>C. R.</given-names>
</name>
<name>
<surname>Gaudreault</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Wei</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Deprez</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Sulea</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Solvated interaction energy: from small-molecule to antibody drug design</article-title>. <source>Front. Mol. Biosci.</source> <volume>10</volume>, <fpage>1210576</fpage>. <pub-id pub-id-type="doi">10.3389/fmolb.2023.1210576</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="book">
<collab>R Development Core Team</collab> (<year>2011</year>). <source>R: a language and environment for statistical computing</source>. <publisher-loc>Vienna, Austria</publisher-loc>: <publisher-name>The R Foundation for Statistical Computing</publisher-name>.</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Radom</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Paci</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Pluckthun</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Computational modeling of designed ankyrin repeat protein complexes with their targets</article-title>. <source>J. Mol. Biol.</source> <volume>431</volume>, <fpage>2852</fpage>&#x2013;<lpage>2868</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmb.2019.05.005</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rothenberger</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Hurdiss</surname>
<given-names>D. L.</given-names>
</name>
<name>
<surname>Walser</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Malvezzi</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Mayor</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Ryter</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>The trispecific DARPin ensovibep inhibits diverse SARS-CoV-2 variants</article-title>. <source>Nat. Biotechnol.</source> <volume>40</volume>, <fpage>1845</fpage>&#x2013;<lpage>1854</lpage>. <pub-id pub-id-type="doi">10.1038/s41587-022-01382-3</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ruffolo</surname>
<given-names>J. A.</given-names>
</name>
<name>
<surname>Chu</surname>
<given-names>L.-S.</given-names>
</name>
<name>
<surname>Mahajan</surname>
<given-names>S. P.</given-names>
</name>
<name>
<surname>Gray</surname>
<given-names>J. J.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Fast, accurate antibody structure prediction from deep learning on massive set of natural antibodies</article-title>, <comment>2004.2020.488972. bioRxiv</comment>. <pub-id pub-id-type="doi">10.1101/2022.04.20.488972</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schilling</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jost</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Ilie</surname>
<given-names>I. M.</given-names>
</name>
<name>
<surname>Schnabl</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Buechi</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Eapen</surname>
<given-names>R. S.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Thermostable designed ankyrin repeat proteins (DARPins) as building blocks for innovative drugs</article-title>. <source>J. Biol. Chem.</source> <volume>298</volume>, <fpage>101403</fpage>. <pub-id pub-id-type="doi">10.1016/j.jbc.2021.101403</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schilling</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Schoppe</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Pluckthun</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2014a</year>). <article-title>From DARPins to LoopDARPins: novel LoopDARPin design allows the selection of low picomolar binders in a single round of ribosome display</article-title>. <source>J. Mol. Biol.</source> <volume>426</volume>, <fpage>691</fpage>&#x2013;<lpage>721</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmb.2013.10.026</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schilling</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Schoppe</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Sauer</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Pluckthun</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2014b</year>). <article-title>Co-crystallization with conformation-specific designed ankyrin repeat proteins explains the conformational flexibility of BCL-W</article-title>. <source>J. Mol. Biol.</source> <volume>426</volume>, <fpage>2346</fpage>&#x2013;<lpage>2362</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmb.2014.04.010</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schneider</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Buchanan</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Taddese</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Deane</surname>
<given-names>C. M.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Dlab: deep learning methods for structure-based virtual screening of antibodies</article-title>. <source>Bioinformatics</source> <volume>38</volume>, <fpage>377</fpage>&#x2013;<lpage>383</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btab660</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Schrag</surname>
<given-names>J. D.</given-names>
</name>
<name>
<surname>Picard</surname>
<given-names>M. E.</given-names>
</name>
<name>
<surname>Gaudreault</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Gagnon</surname>
<given-names>L. P.</given-names>
</name>
<name>
<surname>Baardsnes</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Manenda</surname>
<given-names>M. S.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Binding symmetry and surface flexibility mediate antibody self-association</article-title>. <source>MAbs</source> <volume>11</volume>, <fpage>1300</fpage>&#x2013;<lpage>1318</lpage>. <pub-id pub-id-type="doi">10.1080/19420862.2019.1632114</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sheridan</surname>
<given-names>R. P.</given-names>
</name>
<name>
<surname>McGaughey</surname>
<given-names>G. B.</given-names>
</name>
<name>
<surname>Cornell</surname>
<given-names>W. D.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Multiple protein structures and multiple ligands: effects on the apparent goodness of virtual screening results</article-title>. <source>J. Comput. Aided Mol. Des.</source> <volume>22</volume>, <fpage>257</fpage>&#x2013;<lpage>265</lpage>. <pub-id pub-id-type="doi">10.1007/s10822-008-9168-9</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Strittmatter</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Bertschi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Scheller</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Freitag</surname>
<given-names>P. C.</given-names>
</name>
<name>
<surname>Ray</surname>
<given-names>P. G.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Programmable DARPin-based receptors for the detection of thrombotic markers</article-title>. <source>Nat. Chem. Biol.</source> <volume>18</volume>, <fpage>1125</fpage>&#x2013;<lpage>1134</lpage>. <pub-id pub-id-type="doi">10.1038/s41589-022-01095-3</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Warszawski</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Borenstein Katz</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Lipsh</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Khmelnitsky</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Ben Nissan</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Javitt</surname>
<given-names>G.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Optimizing antibody affinity and stability by the automated design of the variable light-heavy chain interfaces</article-title>. <source>PLoS Comput. Biol.</source> <volume>15</volume>, <fpage>e1007207</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pcbi.1007207</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wood</surname>
<given-names>T. K.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Concerns with computational protein engineering programmes IPRO and OptMAVEn and metabolic pathway engineering programme optStoic</article-title>. <source>Open Biol.</source> <volume>11</volume>, <fpage>200173</fpage>. <pub-id pub-id-type="doi">10.1098/rsob.200173</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yin</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Feng</surname>
<given-names>B. Y.</given-names>
</name>
<name>
<surname>Varshney</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Pierce</surname>
<given-names>B. G.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Benchmarking AlphaFold for protein complex modeling reveals accuracy determinants</article-title>. <source>Protein Sci.</source> <volume>31</volume>, <fpage>e4379</fpage>. <pub-id pub-id-type="doi">10.1002/pro.4379</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Youn</surname>
<given-names>S. J.</given-names>
</name>
<name>
<surname>Kwon</surname>
<given-names>N. Y.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>J. H.</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>J. H.</given-names>
</name>
<name>
<surname>Choi</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>H.</given-names>
</name>
<etal/>
</person-group> (<year>2017</year>). <article-title>Construction of novel repeat proteins with rigid and predictable structures using a shared helix method</article-title>. <source>Sci. Rep.</source> <volume>7</volume>, <fpage>2595</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-017-02803-z</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yu</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Using ggtree to visualize data on tree-like structures</article-title>. <source>Curr. Protoc. Bioinforma.</source> <volume>69</volume>, <fpage>e96</fpage>. <pub-id pub-id-type="doi">10.1002/cpbi.96</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>