<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Archiving and Interchange DTD v2.3 20070202//EN" "archivearticle.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="methods-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Immunol.</journal-id>
<journal-title>Frontiers in Immunology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Immunol.</abbrev-journal-title>
<issn pub-type="epub">1664-3224</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fimmu.2017.00420</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Immunology</subject>
<subj-group>
<subject>Methods</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Large Diversity of Functional Nanobodies from a Camelid Immune Library Revealed by an Alternative Analysis of Next-Generation Sequencing Data</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name><surname>Deschaght</surname> <given-names>Pieter</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x02021;</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Vint&#x000E9;m</surname> <given-names>Ana Paula</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x02021;</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Logghe</surname> <given-names>Marc</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Conde</surname> <given-names>Miguel</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Felix</surname> <given-names>David</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://frontiersin.org/people/u/422362"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Mensink</surname> <given-names>Rob</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Gon&#x000E7;alves</surname> <given-names>Juliana</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Audiens</surname> <given-names>Jorn</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Bruynooghe</surname> <given-names>Yanik</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Figueiredo</surname> <given-names>Rita</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Ramos</surname> <given-names>Diana</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Tanghe</surname> <given-names>Robbe</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Teixeira</surname> <given-names>Daniela</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Van de Ven</surname> <given-names>Liesbeth</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://frontiersin.org/people/u/422579"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Stortelers</surname> <given-names>Catelijne</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Dombrecht</surname> <given-names>Bruno</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="cor1">&#x0002A;</xref>
<uri xlink:href="http://frontiersin.org/people/u/382276"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Ablynx N.V.</institution>, <addr-line>Ghent</addr-line>, <country>Belgium</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: Colin Roger MacKenzie, National Research Council Canada, Canada</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Daniel Zabetakis, United States Naval Research Laboratory, USA; Bryan Briney, Scripps Research Institute, USA; Xin Ge, University of California Riverside, USA</p></fn>
<corresp content-type="corresp" id="cor1">&#x0002A;Correspondence: Bruno Dombrecht, <email>bruno.dombrecht&#x00040;ablynx.com</email></corresp>
<fn fn-type="present-address" id="fn001"><p><sup>&#x02020;</sup>Present address: David Felix, Merus N.V., Utrecht, Netherlands; Rob Mensink, GenCore Facility, Institute for Research and Innovation in Health, University of Porto, Porto, Portugal; Juliana Gon&#x000E7;alves, Fair Journey Biologics, Porto, Portugal; Rita Figueiredo, Immunocore Ltd., Abingdon, UK; Diana Ramos, Fair Journey Biologics, Porto, Portugal; Daniela Teixeira, Fair Journey Biologics, Porto, Portugal; Liesbeth Van de Ven, Argen-X N.V., Ghent, Belgium</p></fn>
<fn fn-type="present-address" id="fn002"><p><sup>&#x02021;</sup>These authors have contributed equally to this work.</p></fn>
<fn fn-type="other" id="fn003"><p>Specialty section: This article was submitted to Vaccines and Molecular Therapeutics, a section of the journal Frontiers in Immunology</p></fn>
</author-notes>
<pub-date pub-type="epub">
<day>10</day>
<month>04</month>
<year>2017</year>
</pub-date>
<pub-date pub-type="collection">
<year>2017</year>
</pub-date>
<volume>8</volume>
<elocation-id>420</elocation-id>
<history>
<date date-type="received">
<day>22</day>
<month>12</month>
<year>2016</year>
</date>
<date date-type="accepted">
<day>24</day>
<month>03</month>
<year>2017</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x000A9; 2017 Deschaght, Vint&#x000E9;m, Logghe, Conde, Felix, Mensink, Gon&#x000E7;alves, Audiens, Bruynooghe, Figueiredo, Ramos, Tanghe, Teixeira, Van de Ven, Stortelers and Dombrecht.</copyright-statement>
<copyright-year>2017</copyright-year>
<copyright-holder>Deschaght, Vint&#x000E9;m, Logghe, Conde, Felix, Mensink, Gon&#x000E7;alves, Audiens, Bruynooghe, Figueiredo, Ramos, Tanghe, Teixeira, Van de Ven, Stortelers and Dombrecht</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) or licensor are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<p>Next-generation sequencing (NGS) has been applied successfully to the field of therapeutic antibody discovery, often outperforming conventional screening campaigns which tend to identify only the more abundant selective antibody sequences. We used NGS to mine the functional nanobody repertoire from a phage-displayed camelid immune library directed to the recepteur d&#x02019;origine nantais (RON) receptor kinase. Challenges to this application of NGS include accurate removal of read errors, correct identification of related sequences, and establishing meaningful inclusion criteria for sequences-of-interest. To this end, a sequence identity threshold was defined to separate unrelated full-length sequence clusters by exploring a large diverse set of publicly available nanobody sequences. When combined with majority-rule consensus building, applying this elegant clustering approach to the NGS data set revealed a wealth of &#x0003E;5,000-enriched candidate RON binders. The huge binding potential predicted by the NGS approach was explored through a set of randomly selected candidates: 90% were confirmed as RON binders, 50% of which functionally blocked RON in an ERK phosphorylation assay. Additional validation came from the correct prediction of all 35 RON binding nanobodies which were identified by a conventional screening campaign of the same immune library. More detailed characterization of a subset of RON binders revealed excellent functional potencies and a promising epitope diversity. In summary, our approach exposes the functional diversity and quality of the outbred camelid heavy chain-only immune response and confirms the power of NGS to identify large numbers of promising nanobodies.</p>
</abstract>
<kwd-group>
<kwd>next-generation sequencing</kwd>
<kwd>clustering</kwd>
<kwd>nanobodies</kwd>
<kwd>recepteur d&#x02019;origine nantais signaling</kwd>
<kwd>phage display</kwd>
<kwd>sequence homology</kwd>
<kwd>amino acid</kwd>
<kwd>immune repertoire diversity</kwd>
</kwd-group>
<counts>
<fig-count count="7"/>
<table-count count="3"/>
<equation-count count="0"/>
<ref-count count="36"/>
<page-count count="11"/>
<word-count count="7508"/>
</counts>
</article-meta>
</front>
<body>
<sec id="S1" sec-type="introduction">
<title>Introduction</title>
<p>Nanobodies are antibody-derived therapeutic proteins based on immunoglobulin single variable domains (<xref ref-type="bibr" rid="B1">1</xref>) derived from the variable domains (VHH) of heavy chain-only antibodies that naturally occur in camelids (<xref ref-type="bibr" rid="B2">2</xref>). Conventionally, nanobodies with desired functional properties are selected from immune, na&#x000EF;ve, or synthetic libraries <italic>via</italic> phage display on the antigen-of-interest (<xref ref-type="bibr" rid="B3">3</xref>). More recently, nanobody libraries have been explored by ribosomal, bacterial, or yeast surface display and by bacterial or yeast two-hybrid selections (<xref ref-type="bibr" rid="B4">4</xref>&#x02013;<xref ref-type="bibr" rid="B10">10</xref>). At the end of this selection process, enriched clones are screened <italic>in vitro</italic> after which hit candidates are identified by means of Sanger sequencing. Although this procedure has a proven track record, the conventional screening approach is often limited to throughputs of several hundreds of clones and thus likely represents only a fraction of the functional potential present in the libraries.</p>
<p>Next-generation sequencing (NGS) technologies have significantly contributed to our knowledge of antibody repertoire diversity in different species or diseases (<xref ref-type="bibr" rid="B11">11</xref>&#x02013;<xref ref-type="bibr" rid="B13">13</xref>). More so, NGS can be a powerful tool in the discovery process of antibody-based therapeutics. The large number of sequencing reads obtained by NGS not only enables unparalleled library quality control but can be applied to more completely assess the binding potential of antibody and nanobody repertoires (<xref ref-type="bibr" rid="B14">14</xref>&#x02013;<xref ref-type="bibr" rid="B21">21</xref>). During the library selection process on the antigen-of-interest, the selective binders are enriched over the background of non-selective clones. A sequence-based frequency analysis then enables the identification of candidate binders which are enriched on the antigen-of-interest in comparison to a negative control condition.</p>
<p>Recepteur d&#x02019;origine nantais (RON) is a receptor tyrosine kinase member of the MET proto-oncogene family (<xref ref-type="bibr" rid="B22">22</xref>, <xref ref-type="bibr" rid="B23">23</xref>). RON dimerization on the cell-surface is required for activation after conformational changes induced by the ligand macrophage-stimulating protein (MSP). Overexpression and splicing variants of RON are implicated in many processes related to cancer initiation, progression, and malignant conversion. Constitutive receptor activation triggers downstream signaling cascades critical for tumorigenesis, including RAS&#x02013;MAPK and PI-3K&#x02013;AKT pathways (<xref ref-type="bibr" rid="B24">24</xref>).</p>
<p>We used NGS to mine a camelid&#x02019;s nanobody selective immune response to human RON (hRON) in comparison to a conventional screening campaign exploring the same immune library for hRON-specific nanobodies. To this end, samples from phage display selections on hRON were sequenced by Illumina MiSeq (2&#x02009;&#x000D7;&#x02009;250&#x02009;bp) which allows for a full coverage of the nanobody encoding sequences. A sequence identity-based clustering approach combined with majority-rule consensus building was utilized, which was developed using publicly available nanobody sequence data. This approach elegantly addressed known issues of PCR and sequencing errors as well as sequence diversity reduction and revealed a wealth of candidate hRON-binding nanobodies. Validation of the method came from the confirmation of all leads which were identified by the conventional screening campaign. In addition, many more functional leads were identified.</p>
</sec>
<sec id="S2" sec-type="materials|methods">
<title>Materials and Methods</title>
<sec id="S2-1">
<title>Proteins, Antibodies, and Cell Lines</title>
<p>Recombinant extracellular domain of human RON (rhRON), and the ligand MSP were purchased from R&#x00026;D Systems (MN, USA). Anti-FLAG antibodies and extravidin peroxidase were purchased from Sigma-Aldrich (MO, USA), goat anti-mouse antibody PE or APC conjugated from Jackson Immuno Research (PA, USA), and anti-M13 monoclonal HRP Conjugate from GE Healthcare. HEK293T (DSMZ, Germany) and llama navel cord fibroblast (Llana) (Ablynx, Belgium) cell lines were transiently transfected using FuGENE HD (Promega, WI, USA) transfection reagent with full-length hRON DNA cloned into pcDNA3.1. The human breast cancer cell line T-47D endogenously expressing RON was obtained from ATCC (VA, USA).</p>
</sec>
<sec id="S2-2">
<title>Immunizations, Library Construction, and Phage Display Selections</title>
<p>Recepteur d&#x02019;origine nantais-targeting nanobodies were generated through immunization of a llama with rhRON, essentially as described elsewhere (<xref ref-type="bibr" rid="B3">3</xref>). Briefly, a llama was immunized first with 100&#x02009;&#x000B5;g of protein followed by three times 50&#x02009;&#x000B5;g, after which blood samples were taken. Phage display libraries derived from peripheral blood mononuclear cells (PBMCs) were prepared and used as previously described (<xref ref-type="bibr" rid="B3">3</xref>). The VHH fragments were cloned into a M13 phagemid vector containing the FLAG<sub>3</sub> and His<sub>6</sub> tags. The resulting library size was 4.8&#x02009;&#x000D7;&#x02009;10<sup>8</sup> with 91% of insert. The library was rescued by infecting exponentially growing <italic>Escherichia coli</italic> TG1 [(F&#x02032; <italic>traD36 proAB lacIqZ</italic> &#x00394;<italic>M15</italic>) <italic>supE thi-1</italic> &#x00394;<italic>(lac-proAB)</italic> &#x00394;<italic>(mcrB-hsdSM)5(rK&#x02212; mK&#x02212;)</italic>] cells followed by superinfection with VCSM13 helper phage, resulting in 4.4&#x02009;&#x000D7;&#x02009;10<sup>13</sup> cfu/ml. For the NGS samples, the RON and the negative control outputs, with sizes of respectively, 8&#x02009;&#x000D7;&#x02009;10<sup>6</sup> and 9&#x02009;&#x000D7;&#x02009;10<sup>5</sup> cfu, were derived from one round of selection on HEK293T cells expressing hRON and on HEK293T cells, respectively. For the conventional screening campaign, phage display selections were performed on HEK239T or Llana cells expressing hRON and on rhRON protein either directly immobilized on plate or captured <italic>via</italic> biotin by streptavidin-coated magnetic beads (Dynabeads, Invitrogen). The phage outputs were rescued as described above for the library. For screening purposes, <italic>E. coli</italic> TG1 cells were infected with the resulting phage outputs and individual colonies were grown in 96-deep-well plates. The expression of monoclonal nanobodies was induced by addition of IPTG and the crude periplasmic extracts containing the nanobodies were prepared by freeze-thawing of the bacterial pellets overnight in PBS followed by centrifugation to remove cell debris.</p>
</sec>
<sec id="S2-3">
<title>Cloning and Production of Nanobodies</title>
<p>Synthetic DNA fragments (Integrated DNA Technologies, Belgium) encoding nanobodies from the NGS campaign and nanobody genes derived from the conventional screening approach were cloned into an expression vector in frame with an N-terminal OmpA signal peptide and C-terminal FLAG<sub>3</sub> and His<sub>6</sub> tags. Production and purification were in essence performed as described before (<xref ref-type="bibr" rid="B3">3</xref>).</p>
</sec>
<sec id="S2-4">
<title>NGS Sample Preparation and Sequencing</title>
<p>Polyclonal plasmid DNA preparations from <italic>E. coli</italic> cultures infected with two different phage samples (RON and negative control) were used as PCR template. The first PCR was performed with primers FR1 (5&#x02032;-GAGGTGCAGCTGGTGGAGTCT-3&#x02032;, encoding EVQLVES) and FR4 (5&#x02032;-TGAGGAGACGGTGACCWGGGT-3&#x02032;, encoding T(L/Q)VTVSS). For each sample, 48 parallel PCR reactions were run with KAPA HiFi DNA polymerase (Kapa Biosystems) using the following protocol: 3&#x02009;min at 95&#x000B0;C; 20 cycles of 20&#x02009;s at 98&#x000B0;C, 25&#x02009;s at 55&#x000B0;C, 10&#x02009;s at 72&#x000B0;C; once 5&#x02009;min at 72&#x000B0;C. After PCR all samples were subjected to sample clean-up (PureLink PCR Purification Kit, Life Technologies). In a second PCR, these DNA amplicons were flanked by barcoded i7 TruSeq adapters as prescribed (Illumina). The samples were sequenced on a MiSeq system using the Illumina v2 2&#x02009;&#x000D7;&#x02009;250&#x02009;bp chemistry kit.</p>
</sec>
<sec id="S2-5">
<title>NGS Data Processing</title>
<p>In a first step, the reads were sorted by barcode, followed by barcode and Illumina TruSeq adapter clipping with bcl2fastq 1.8.4 (Illumina). Forward and reverse reads were combined using open source software FLASH 1.2.4 (<xref ref-type="bibr" rid="B25">25</xref>) available from <uri xlink:href="https://ccb.jhu.edu/software/FLASH/">https://ccb.jhu.edu/software/FLASH/</uri> (minimal overlap: 10 bases, maximum mismatch rate: 25%). After PCR primer sequence detection, the reads were turned into the forward (FR1 primer) to reverse (FR4 primer) orientation and reads with average Phred scores &#x0003C;38 or lengths &#x0003C;150&#x02009;bp were discarded. After translation with the freely available BioPhyton package<xref ref-type="fn" rid="fn1"><sup>1</sup></xref> in frame &#x0002B;1, starting at the 5&#x02032;-end of the FR1 PCR primer, peptides ending in frame with the FR4 primer sequence were considered valid, thus excluding reads with frameshifts and/or premature stop codons.</p>
</sec>
<sec id="S2-6">
<title>Downloading Nanobody Sequences</title>
<p>Publicly available nanobody sequences were downloaded from the NCBI Protein database<xref ref-type="fn" rid="fn2"><sup>2</sup></xref> (accessed 15 March 2016) using &#x0201C;<italic>(camelidae[Organism]) AND ((VHH) OR (Nanobody) OR (single domain)) AND immunoglobulin</italic>&#x0201D; as query. Nine hundred forty-five matches were obtained and aligned. After visual inspection of this alignment, obvious non-nanobody and truncated or partial sequences were manually removed, leaving 888 sequences for clustering (Table S1 in Supplementary Material).</p>
</sec>
<sec id="S2-7">
<title>Nanobody Clustering and Alignment</title>
<p>Before clustering, the residues corresponding to IMGT V-DOMAIN positions 1&#x02013;7 and 122&#x02013;128, the first and last seven residues of FR1 and FR4, respectively (<xref ref-type="bibr" rid="B26">26</xref>), were trimmed from the nanobody peptide sequences. This was done in order to remove undesirable sequence variation introduced by the PCR primers used in the preparation of the NGS samples or coming from partial FR1 and/or FR4 regions in publicly available nanobody sequences. The trimmed peptide sequences were clustered with CD-HIT version 4.6.1 (<xref ref-type="bibr" rid="B27">27</xref>, <xref ref-type="bibr" rid="B28">28</xref>). A detailed user manual as well as a web server of this freely available and widely used clustering software package can be found at <uri xlink:href="http://weizhongli-lab.org/cd-hit/">http://weizhongli-lab.org/cd-hit/</uri>. The program was run in the slow/accurate mode (&#x02212;<italic>g</italic>&#x02009;&#x0003D;&#x02009;1), no length differences were allowed (length difference cutoff &#x02212;<italic>s</italic>&#x02009;&#x0003D;&#x02009;1), and different identity cutoffs (sequence identity threshold &#x02212;<italic>c</italic>&#x02009;&#x0003D;&#x02009;0.70, 0.75, 0.80, 0.85, 0.90, 0.95, and 1.00) were evaluated.</p>
<p>Alignments were generated with CLC Main Workbench version 7.6.4 (Qiagen).</p>
</sec>
<sec id="S2-8">
<title>Binding ELISA</title>
<p>rhRON (1&#x02009;&#x000B5;g/ml) was immobilized directly on 384-well microtiter plates. Free-binding sites were blocked by 4% Marvel in PBS. Next, 5&#x02009;&#x000B5;l of crude periplasmic extracts in 50&#x02009;&#x000B5;l 2% Marvel PBST were added. Nanobody binding was revealed using a mouse-anti-FLAG HRP-conjugated antibody. The OD<sub>450nm</sub> values of each clone were divided by those of a negative control nanobody and considered positive if the resulting ratio was &#x02265;2.</p>
</sec>
<sec id="S2-9">
<title>Epitope Binning</title>
<p>Biotinylated rhRON (1&#x02009;nM) was captured by NeutrAvidin immobilized on 96-well microtiter plates (2&#x02009;&#x000B5;g/ml) and blocked by 1% casein in PBS. Next, 1&#x02009;&#x000B5;l of purified monoclonal phage (10<sup>11</sup> cfu/ml) displaying nanobody in 100&#x02009;&#x000B5;l, 0.1% casein PBST were added in the presence and absence of crude periplasmic extract containing nanobodies at 1/10 dilutions. Phage binding was detected <italic>via</italic> anti-M13 HRP-conjugated antibody. Competition for binding to an overlapping epitope was revealed by the drop in signal of phage binding in the presence of the nanobody in the crude periplasmic extract.</p>
</sec>
<sec id="S2-10">
<title>Off-rate Determination</title>
<p>Off-rates were determined by surface plasmon resonance of crude periplasmic extracts on a ProteOn instrument (Biorad, CA, USA). rhRON was immobilized to GLC sensor chips surface and nanobody binding was assessed using 1/10 diluted periplasmic extracts. Each nanobody was injected for 2&#x02009;min at a flow rate of 45&#x02009;&#x000B5;l/min to allow binding to chip-bound antigen. Next, binding buffer without nanobody was injected at the same flow rate to allow spontaneous dissociation of bound nanobody. Regeneration was done with 10&#x02009;mM glycine HCl, pH2.5. From the sensorgrams obtained for the different nanobodies <italic>k</italic><sub>off</sub> values were calculated. Data processing and analysis were done with the ProteOn Manager Software, Version 2.1.1.18 applying the Langmuir kinetic model.</p>
</sec>
<sec id="S2-11">
<title>Inhibition of MSP-Induced ERK Phosphorylation</title>
<p>Functional blockade of RON kinase activation by nanobodies was assessed by inhibition of ligand-induced MAPK activation in T-47D breast cancer cells. For screening purposes, the AlphaLISA SureFire Ultra phospho-ERK 1/2 (Thr202/Tyr204) kit was used (PerkinElmer, MA, USA). T-47D cells (2.0&#x02009;&#x000D7;&#x02009;10<sup>4</sup>/well in 0.1&#x02009;ml) were seeded in 96-wells plates in culture medium, incubated for 24&#x02009;h after which the medium was replaced by serum-free medium to synchronize the cells overnight. Cells were pre-incubated with nanobodies present in crude periplasmic extract (1/25 dilution) for 1&#x02009;h, after which the RON receptor was stimulated by addition of 3.5&#x02009;nM of MSP for 15&#x02009;min at 37&#x000B0;C. The cells were resuspended in 60&#x02009;&#x000B5;l of lysis buffer after removal of the medium. The amount of phosphorylated ERK versus total ERK was determined following the recommendations from the provider. Inhibition % was calculated using non-stimulated cells and crude periplasmic extract of irrelevant control nanobody as references. For IC<sub>50</sub> determination, serum-starved T-47D cells (3.5&#x02009;&#x000D7;&#x02009;10<sup>4</sup> cells/well) were incubated with serial dilutions of purified nanobodies (duplicates, starting at 0.6&#x02009;&#x000B5;M) and stimulated with 1&#x02009;nM of MSP for 30&#x02009;min at 37&#x000B0;C. Quantification of the cellular the pErk levels was done using the HTRF phospho-ERK (Thr202/Tyr204) Assay (Cisbio, France).</p>
</sec>
<sec id="S2-12">
<title>Cell Binding Assays</title>
<p>To screen for binding to cell-expressed RON, 1/10 diluted crude periplasmic extracts were incubated with HEK293T-hRON and HEK293T cells (5&#x02009;&#x000D7;&#x02009;10<sup>4</sup> cells/well) in FACS buffer (PBS supplemented with 10% fetal bovine serum and 0.05% sodium azide). Nanobody binding was detected using mouse anti-FLAG antibodies followed by goat anti-mouse APC conjugate. Mean fluorescence intensity values of each clone on the HEK293T-hRON cells were divided by those of the background signal, normalized to the same ratio on HEK293T cells. Clones were considered positive with a ratio &#x02265;2. To determine EC<sub>50</sub> values, a dilution series of purified nanobodies starting at 500&#x02009;nM in duplicates was added to T-47D cells (1&#x02009;&#x000D7;&#x02009;10<sup>5</sup>/well). The detection was carried out as described above.</p>
</sec>
<sec id="S2-13">
<title>Ligand Competition ELISA</title>
<p>A competition ELISA was used to determine blockade of the binding of the MSP ligand to rhRON. rhRON (1&#x02009;&#x000B5;g/ml) was immobilized directly on 96-well microtiter plates. Free-binding sites were blocked using 4% Marvel in PBS for 1&#x02009;h at room temperature. Next, a dilution series of purified nanobodies starting at 1&#x02009;&#x000B5;M (in duplicate) was added simultaneously with 2&#x02009;nM in-house biotinylated MSP in 100&#x02009;&#x000B5;l 2% Marvel PBST. MSP binding was detected <italic>via</italic> extravidin peroxidase.</p>
</sec>
<sec id="S2-14">
<title>Calculations</title>
<p>The sequence counts per cluster in the RON sample were multiplied with a factor of 1.21 (3.4&#x02009;&#x000D7;&#x02009;10<sup>6</sup>/2.8&#x02009;&#x000D7;&#x02009;10<sup>6</sup>) to normalize for the difference in total counts with the negative control sample (Table <xref ref-type="table" rid="T1">1</xref>). The enrichment factor of a cluster was calculated as follows: number of sequences (normalized counts) in the RON sample belonging to that cluster divided by the number of sequences (counts) in the negative control sample belonging to the same cluster. For clusters present in the RON sample but not in the negative control sample, the counts in the latter were changed from 0 to 1 in order to calculate the enrichment factor.</p>
<table-wrap position="float" id="T1">
<label>Table 1</label>
<caption><p><bold>Summary of next-generation sequencing raw data and initial processing output</bold>.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th valign="top" align="center"/>
<th valign="top" align="center">Negative control</th>
<th valign="top" align="center">Recepteur d&#x02019;origine nantais</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">Selection output size (cfu)</td>
<td align="center" valign="top">9&#x02009;&#x000D7;&#x02009;10<sup>5</sup></td>
<td align="center" valign="top">8&#x02009;&#x000D7;&#x02009;10<sup>6</sup></td>
</tr>
<tr>
<td align="left" valign="top">Raw reads (counts)</td>
<td align="center" valign="top">1.0&#x02009;&#x000D7;&#x02009;10<sup>7</sup></td>
<td align="center" valign="top">7.5&#x02009;&#x000D7;&#x02009;10<sup>6</sup></td>
</tr>
<tr>
<td align="left" valign="top">Joined reads (counts)</td>
<td align="center" valign="top">4.9&#x02009;&#x000D7;&#x02009;10<sup>6</sup></td>
<td align="center" valign="top">3.6&#x02009;&#x000D7;&#x02009;10<sup>6</sup></td>
</tr>
<tr>
<td align="left" valign="top">Joinable fraction (%)</td>
<td align="center" valign="top">94</td>
<td align="center" valign="top">96</td>
</tr>
<tr>
<td align="left" valign="top">Full-length nanobody sequences (counts)</td>
<td align="center" valign="top">3.4&#x02009;&#x000D7;&#x02009;10<sup>6</sup></td>
<td align="center" valign="top">2.8&#x02009;&#x000D7;&#x02009;10<sup>6</sup></td>
</tr>
<tr>
<td align="left" valign="top">Unique sequences (counts)</td>
<td align="center" valign="top">1.8&#x02009;&#x000D7;&#x02009;10<sup>6</sup></td>
<td align="center" valign="top">1.1&#x02009;&#x000D7;&#x02009;10<sup>6</sup></td>
</tr>
<tr>
<td align="left" valign="top">Fraction unique sequences (%)</td>
<td align="center" valign="top">53</td>
<td align="center" valign="top">39</td>
</tr>
<tr>
<td align="left" valign="top">Unique sequences/selection output size (%)</td>
<td align="center" valign="top">200</td>
<td align="center" valign="top">14</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>Confidence intervals of proportions, EC<sub>50</sub>, and IC<sub>50</sub> values were calculated with GraphPad Prism 6 (GraphPad Software).</p>
</sec>
</sec>
<sec id="S3">
<title>Results</title>
<p>Next-generation sequencing was used to mine the functional nanobody repertoire from a camelid immune library. The experiments below describe the NGS-based approach to identify RON-selective nanobodies from an immune library in comparison with a conventional screening campaign (Figure <xref ref-type="fig" rid="F1">1</xref>).</p>
<fig id="F1" position="float">
<label>Figure 1</label>
<caption><p><bold>Schematic overview of the work flows for the next-generation sequencing and conventional screening campaigns</bold>.</p></caption>
<graphic xlink:href="fimmu-08-00420-g001.tif"/>
</fig>
<sec id="S3-1">
<title>NGS: Raw Data Processing</title>
<p>A nanobody phage library was constructed from the PBMCs obtained from a llama immunized with rhRON. The phages were subjected for one selection round to HEK293T cells overexpressing hRON or to the parental HEK293T cells acting as negative control. The nanobody sequences were PCR-amplified from the resulting outputs, introducing a different DNA barcode to each sample (negative control and RON). A total of 1.75&#x02009;&#x000D7;&#x02009;10<sup>7</sup> raw reads were obtained (MiSeq Kit v2, 2&#x02009;&#x000D7;&#x02009;250&#x02009;bp). After barcode deconvolution and clipping, 95% of the forward reads could be joined to their corresponding reverse reads. Translation of the joined DNA reads excluded 23&#x02013;30% of the reads for further analysis caused by the introduction of frameshifts and/or premature stop codons. These clean-up steps yielded around 3&#x02009;&#x000D7;&#x02009;10<sup>6</sup> full-length nanobody sequences per sample (Table <xref ref-type="table" rid="T1">1</xref>). A difference between the samples was observed with respect to sequence diversity: the negative control sample contained relatively more unique sequences, compared to the RON sample (Table <xref ref-type="table" rid="T1">1</xref>). Consistent with published data (<xref ref-type="bibr" rid="B14">14</xref>, <xref ref-type="bibr" rid="B15">15</xref>), this suggests that the selection process enriched for RON binders, resulting in a reduction of overall sequence diversity. The selection output sizes (Table <xref ref-type="table" rid="T1">1</xref>) represent the maximum possible sequence diversity of the sequenced samples. Strong amplification of identical binders by the phage display process explains the 14% ratio of unique sequences over selection output size in the RON sample (Table <xref ref-type="table" rid="T1">1</xref>). The observation that the negative control sample appears to have twofold (200% ratio) more unique sequences than theoretically possible, can be explained as follows. First, it is reasonable to assume a twofold error on the quantification of the selection output size, which was done by titrating out a phage-infected <italic>E. coli</italic> culture, followed by a count of colony forming units (cfu). Secondly, the different downstream PCR amplification steps and the actual MiSeq sequencing will have introduced errors resulting in an increased diversity.</p>
</sec>
<sec id="S3-2">
<title>NGS: Nanobody Sequence Clustering and Frequency Analysis</title>
<p>The large number of unique sequences, &#x0003E;1&#x02009;&#x000D7;&#x02009;10<sup>6</sup> per sample (Table <xref ref-type="table" rid="T1">1</xref>), prompted us to first explore a meaningful reduction of the sequence diversity, before performing an enrichment analysis to identify candidate hRON binders. More so, it is well known that errors introduced by the different PCR and sequencing steps significantly hamper the correct analysis of antibody repertoire sequence diversity, especially for true rare clones (<xref ref-type="bibr" rid="B11">11</xref>). Different methodologies have been explored to address error reduction, including CDR-based clustering or clonotyping, frequency-based consensus building, and replicate sequencing. Clustering of related sequences (clonal grouping and B-cell lineage trees) has been used extensively in the field of antibody repertoire sequencing to meaningfully reduce sequence diversity (<xref ref-type="bibr" rid="B12">12</xref>, <xref ref-type="bibr" rid="B13">13</xref>). However, the main challenge here is to define a sequence identity threshold that allows for the correct clustering of (clonally) related sequences.</p>
<p>To this purpose, it was decided to explore sequence diversity and relatedness in a large set of publicly available nanobody sequences. Nanobody sequences were downloaded, curated (see <xref ref-type="sec" rid="S2">Materials and Methods</xref>) and the 888 sequences thus obtained were further reduced to a non-redundant set of 629 unique nanobodies. Sequence clustering was done using CD-HIT (<xref ref-type="bibr" rid="B27">27</xref>, <xref ref-type="bibr" rid="B28">28</xref>), a freely available program to efficiently handle extremely large datasets. Briefly, the algorithm sorts input sequences from long to short and processes them sequentially. The first sequence is classified as the first cluster representative, after which each of the remaining sequences is compared to the representative sequences found before it and classified as redundant or representative based on similarity. The public dataset was clustered at sequence identity thresholds ranging from 0.7 (70% identity) to 1.0 (100% identity) with no length differences being allowed. For each of the resulting clusters, we checked whether its members were related nanobodies or not. The term related as used here, refers to either targeting the same antigen, originating from the same publication, or sharing a database submission origin (date and authors). Lowering sequence identity thresholds lead to a continuous increase in cluster size (number of sequences per cluster), in number of clusters containing unrelated nanobody sequences, and in number of unrelated nanobody sequences per cluster (Figure <xref ref-type="fig" rid="F2">2</xref>; Table S1 in Supplementary Material). To illustrate this trend better, sequence alignments were generated (Figure S1 in Supplementary Material) with the members of a few representative clusters, identified by capital letters in Figure <xref ref-type="fig" rid="F2">2</xref>. The sequences captured by clusters A, B, and C, respectively, target the same antigen and thus are deemed related. On the other hand, most of the sequences captured by clusters D, E, and F target different antigens and thus are qualified as unrelated. Based on these findings, it was decided to apply an identity threshold of 0.9 for the further analysis of the NGS data set.</p>
<fig id="F2" position="float">
<label>Figure 2</label>
<caption><p><bold>Clustering of publicly available nanobody sequences</bold>. On the <italic>x</italic>-axis, the different CD-HIT clustering exercises at various sequence identity thresholds are shown, including the number of clusters at a given threshold. The <italic>y</italic>-axis (cluster size) displays the sequence counts per cluster. The symbol size indicates the number of unrelated nanobody sequences. The identities of the sequences in each cluster are given in Table S1 in Supplementary Material. The alignments of the sequences captured in clusters identified by a capital letter are shown in Figure S1 in Supplementary Material.</p></caption>
<graphic xlink:href="fimmu-08-00420-g002.tif"/>
</fig>
<p>Similar to the observation with the clusters of unique (100% identical) sequences (Table <xref ref-type="table" rid="T1">1</xref>), the negative control sample had more clusters with a sequence identity threshold of 0.9 than the RON sample (Table <xref ref-type="table" rid="T2">2</xref>). The threefold reduction in number of clusters in the RON sample compared to the negative control sample indicates a decrease in sequence diversity driven by the positive selection pressure. Clusters were subdivided in three groups, based on size: orphan clusters have one single member, medium clusters contain 2&#x02013;10 members, and large clusters contain &#x0003E;10 members. After selection on the antigen, a reduction in number of orphan and medium clusters was observed also here, while the number of large clusters increased (Table <xref ref-type="table" rid="T2">2</xref>), suggestive of positive selection pressure for clusters of hRON-binding sequences. Accordingly, the fraction of sequences present in large clusters and the mean cluster sizes increased after the selection on hRON (Table <xref ref-type="table" rid="T2">2</xref>).</p>
<table-wrap position="float" id="T2">
<label>Table 2</label>
<caption><p><bold>Summary of next-generation sequencing CD-HIT 0.9 clusters</bold>.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th valign="top" align="center"/>
<th valign="top" align="center">Negative control</th>
<th valign="top" align="center">Recepteur d&#x02019;origine nantais</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">All clusters (count)</td>
<td align="center" valign="top">8.1&#x02009;&#x000D7;&#x02009;10<sup>5</sup></td>
<td align="center" valign="top">2.7&#x02009;&#x000D7;&#x02009;10<sup>5</sup></td>
</tr>
<tr>
<td align="left" valign="top">Mean cluster size (&#x00023; sequences)</td>
<td align="center" valign="top">4</td>
<td align="center" valign="top">11</td>
</tr>
<tr>
<td align="left" valign="top">Orphan clusters (1 member) (count)</td>
<td align="center" valign="top">6.5&#x02009;&#x000D7;&#x02009;10<sup>5</sup></td>
<td align="center" valign="top">1.9&#x02009;&#x000D7;&#x02009;10<sup>5</sup></td>
</tr>
<tr>
<td align="left" valign="top">Fraction of total sequences (%)</td>
<td align="center" valign="top">19</td>
<td align="center" valign="top">7</td>
</tr>
<tr>
<td align="left" valign="top">Medium clusters (1&#x02009;&#x0003C;&#x02009;<italic>n</italic>&#x02009;&#x02264;&#x02009;10 members) (count)</td>
<td align="center" valign="top">1.3&#x02009;&#x000D7;&#x02009;10<sup>5</sup></td>
<td align="center" valign="top">6.5&#x02009;&#x000D7;&#x02009;10<sup>4</sup></td>
</tr>
<tr>
<td align="left" valign="top">Fraction of total sequences (%)</td>
<td align="center" valign="top">14</td>
<td align="center" valign="top">8</td>
</tr>
<tr>
<td align="left" valign="top">Mean cluster size (&#x00023; sequences)</td>
<td align="center" valign="top">3.8</td>
<td align="center" valign="top">3.4</td>
</tr>
<tr>
<td align="left" valign="top">Large clusters (<italic>n</italic>&#x02009;&#x0003E;&#x02009;10 members) (count)</td>
<td align="center" valign="top">3.1&#x02009;&#x000D7;&#x02009;10<sup>4</sup></td>
<td align="center" valign="top">1.2&#x02009;&#x000D7;&#x02009;10<sup>4</sup></td>
</tr>
<tr>
<td align="left" valign="top">Fraction of total sequences (%)</td>
<td align="center" valign="top">67</td>
<td align="center" valign="top">86</td>
</tr>
<tr>
<td align="left" valign="top">Mean cluster size (&#x00023; sequences)</td>
<td align="center" valign="top">75</td>
<td align="center" valign="top">208</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>Besides cluster size, also the enrichment factor (ratio of sequence counts per cluster in RON sample over negative control sample) can be considered as a meaningful parameter to select candidate RON-specific nanobodies. To add more statistical robustness to our analysis, only clusters with a size &#x02265;10 and an enrichment factor &#x02265;10 were considered. These inclusion criteria resulted in a &#x0003E;50-fold reduction in the number of clusters from 2.7&#x02009;&#x000D7;&#x02009;10<sup>5</sup> to 5,173 (Table <xref ref-type="table" rid="T2">2</xref>; Figure <xref ref-type="fig" rid="F3">3</xref>). The resulting large panel of 5,173 clusters with a sequence identity threshold of 0.9&#x02014;all different candidate hRON binders&#x02014;has enrichment factors of up to 3,000 and cluster sizes of up to 2.7&#x02009;&#x000D7;&#x02009;10<sup>5</sup> counts (Figure <xref ref-type="fig" rid="F3">3</xref>).</p>
<fig id="F3" position="float">
<label>Figure 3</label>
<caption><p><bold>Next-generation sequencing (NGS) frequency analysis identifies 5,173 candidate human RON binders</bold>. All symbols represent CD-HIT clusters (0.9 sequence identity threshold) with cluster sizes [sequence counts in the recepteur d&#x02019;origine nantais (RON) sample] &#x02265;10 and enrichment factors (ratio of sequence counts per cluster in RON sample over negative control sample) &#x02265;10. Blue squares represent the clusters that were also identified by the conventional screening campaign. Green triangles represent the NGS clusters that were selected for further screening. Clusters for which no sequence counts were observed in the negative control sample were attributed a sequence count of one, in order to be able to calculate and plot enrichment factors for these clusters.</p></caption>
<graphic xlink:href="fimmu-08-00420-g003.tif"/>
</fig>
</sec>
<sec id="S3-3">
<title>Binding and Functional Characterization of RON Nanobodies</title>
<p>In the conventional screening campaign, the same immune phage library was selected for up to two rounds on cells overexpressing hRON and/or on rhRON. Crude periplasmic extracts of enriched single clones were evaluated for binding by ELISA and FACS, followed by Sanger sequencing of the hits (Table S2 in Supplementary Material; blue squares in Figure <xref ref-type="fig" rid="F4">4</xref>A). Sequence analysis revealed that all 35 nanobodies derived from the conventional screening were correctly identified by the NGS approach (blue squares in Figure <xref ref-type="fig" rid="F3">3</xref>) with enrichment factors ranging from 13 to 1,364 and cluster sizes ranging from as low as 16 to as high as 2.7&#x02009;&#x000D7;&#x02009;10<sup>5</sup> counts (Table S2 in Supplementary Material). The conventional screening approach tended to identify the most abundant sequences: 17 out of 22 clusters (77%) with a cluster size &#x0003E;1.0&#x02009;&#x000D7;&#x02009;10<sup>4</sup> counts were also found <italic>via</italic> the conventional screening. However, small clusters with relatively small enrichment factors were also identified by the conventional approach (see blue squares in bottom left quadrant of Figure <xref ref-type="fig" rid="F3">3</xref>). The fact that all 35 conventionally identified nanobodies were captured by the 5,173 NGS clusters validates our frequency-based CD-HIT clustering NGS approach as an efficient method to identify binders. At the same time, it emphasizes the huge binding potential of the immune library that is left untapped by the conventional approach, which in this particular case means that the RON library could theoretically contain &#x0003E;100 times more binders.</p>
<fig id="F4" position="float">
<label>Figure 4</label>
<caption><p><bold>(A)</bold> Binding to human RON (hRON) of candidate binders. Shown are the selective binding ratios of ELISA and FACS experiments. <bold>(B)</bold> Inhibition of ligand-induced ERK phosphorylation by candidate binders. Shown are the % inhibition of ERK phosphorylation and the selective binding ratios of the ELISA experiment. Green triangles represent 28 randomly selected candidate hRON-binding nanobodies, predicted by the next-generation sequencing (NGS) analysis. Blue squares represent 35 hRON-binding nanobodies, predicted by the NGS analysis and identified in the conventional screening campaign. The white triangle represents nanobody NGS00009 which was not analyzed in the FACS experiment (see Table S2 in Supplementary Material) and as such was given a selective binding ratio of 0, but scored positive in the ELISA and pERK assays.</p></caption>
<graphic xlink:href="fimmu-08-00420-g004.tif"/>
</fig>
<p>To explore the untapped binding potential predicted by the NGS analysis, 28 additional clusters were randomly selected for evaluation in hRON-binding ELISA and FACS. The selected clusters represent a range of enrichment factors from 13 to 406 and cluster sizes from 309 to 1.4&#x02009;&#x000D7;&#x02009;10<sup>4</sup> counts (Table S2 in Supplementary Material; green triangles in Figure <xref ref-type="fig" rid="F3">3</xref>). The majority-rule consensus, derived from the alignment of all the sequences that make up a given cluster, was then used as the sequence representative of that cluster. In this manner, the sequence information of the most abundantly present (enriched) sequences in a given cluster is efficiently captured while at the same time PCR and read errors are filtered out (<xref ref-type="bibr" rid="B11">11</xref>). The consensus sequences were reverse translated, ordered as synthetic DNA, and cloned into an <italic>E. coli</italic> expression vector. Crude periplasmic extracts of each clone were used to assess binding to hRON in ELISA and FACS. Of these randomly selected NGS nanobodies, 25/28 (89%, with 95% confidence interval of 72&#x02013;98%) bind to hRON with comparable binding levels to the clones also identified by the conventional campaign (Table S2 in Supplementary Material; compare green triangles to blue squares in Figure <xref ref-type="fig" rid="F4">4</xref>A). Moreover, 14/25 (56%, with 95% confidence interval of 35&#x02013;76%) of the randomly selected binders show functional blockade in the MSP-induced ERK phosphorylation assay (Figure <xref ref-type="fig" rid="F4">4</xref>B).</p>
<p>An interesting observation is that there is no clear correlation between the binding strength of a given cluster&#x02014;as measured by its ELISA ratio to rhRON&#x02014;and its size or enrichment factor (Figure <xref ref-type="fig" rid="F5">5</xref>). In other words, it is probably ill-advised to overly focus on cluster size or enrichment factor as sole inclusion criteria for candidate binders. Good binders can be found in any quadrant of Figure <xref ref-type="fig" rid="F3">3</xref>. Extrapolating from the data of the 28 randomly selected clusters, we speculate that around 90% of the &#x0003E;5,000 remaining unexplored clusters could constitute RON binders, of which more than half could interfere with RON function.</p>
<fig id="F5" position="float">
<label>Figure 5</label>
<caption><p><bold>Absence of correlation between next-generation sequencing cluster size or enrichment factor and binding strength to recepteur d&#x02019;origine nantais (RON)</bold>. Shown are selective binding ratios from the ELISA experiment of each candidate human RON-binding nanobody and the <bold>(A)</bold> size (sequence counts in the RON sample) or <bold>(B)</bold> enrichment factor (ratio of sequence counts per cluster in RON sample over negative control sample) of the corresponding clusters.</p></caption>
<graphic xlink:href="fimmu-08-00420-g005.tif"/>
</fig>
<p>Twelve hRON-binding nanobodies identified by both the NGS and conventional approaches were further characterized as purified protein. Binding affinities were assessed on T-47D cells endogenously expressing RON, indicating EC<sub>50</sub> values ranging from &#x0003E;1&#x02009;&#x003BC;M to 50&#x02009;pM, and with off-rates ranging from 7&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;3</sup> to 3&#x02009;&#x000D7;&#x02009;10<sup>-4</sup> s<sup>&#x02212;1</sup> (Table <xref ref-type="table" rid="T3">3</xref>). More so, all nanobodies completely inhibited MSP-induced ERK phosphorylation in T-47D cells with IC<sub>50</sub> values ranging from 300 to 5&#x02009;nM (Table <xref ref-type="table" rid="T3">3</xref>; Figure <xref ref-type="fig" rid="F6">6</xref>). Nine out of twelve fully block the binding of the MSP ligand to hRON, three others are competing only poorly&#x02014;if at all&#x02014;with MSP binding in the tested concentration range (Table <xref ref-type="table" rid="T3">3</xref>; Figure <xref ref-type="fig" rid="F6">6</xref>). Competition experiments revealed that the nanobodies could be assigned to four non-overlapping epitope bins. Two of the nanobodies share a competing footprint with two epitope bins. Together these data indicate that functionally inhibiting anti-hRON nanobodies are present in the immune repertoire with good potencies and epitope diversity.</p>
<table-wrap position="float" id="T3">
<label>Table 3</label>
<caption><p><bold>Overview characterization of selected anti-human RON nanobodies</bold>.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th valign="top" align="left">ID</th>
<th valign="top" align="center"><italic>k</italic><sub>off</sub> (s<sup>&#x02212;1</sup>)</th>
<th valign="top" align="center">EC<sub>50</sub> (M) binding</th>
<th valign="top" align="center">IC<sub>50</sub> (M) inhibition of MSP binding<xref ref-type="table-fn" rid="tfn1"><sup>a</sup></xref></th>
<th valign="top" align="center">IC<sub>50</sub> (M) inhibition of ERK phosphorylation<xref ref-type="table-fn" rid="tfn1"><sup>a</sup></xref></th>
<th valign="top" align="center">Epitope bin</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">8A09</td>
<td align="center" valign="top">6.2&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;4</sup></td>
<td align="center" valign="top">9.3&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;11</sup></td>
<td align="center" valign="top">1.6&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (98%)</td>
<td align="center" valign="top">4.9&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;9</sup> (100%)</td>
<td align="center" valign="top">A</td>
</tr>
<tr>
<td align="left" valign="top">8F09</td>
<td align="center" valign="top">6.4&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;4</sup></td>
<td align="center" valign="top">2.0&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;10</sup></td>
<td align="center" valign="top">1.2&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (98%)</td>
<td align="center" valign="top">5.2&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;9</sup> (99%)</td>
<td align="center" valign="top">A</td>
</tr>
<tr>
<td align="left" valign="top">11F05</td>
<td align="center" valign="top">4.6&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;4</sup></td>
<td align="center" valign="top">5.3&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;11</sup></td>
<td align="center" valign="top">9.7&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;9</sup> (98%)</td>
<td align="center" valign="top">6.0&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;9</sup> (100%)</td>
<td align="center" valign="top">A</td>
</tr>
<tr>
<td align="left" valign="top">8A12</td>
<td align="center" valign="top">2.8&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;3</sup></td>
<td align="center" valign="top">1.2&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;10</sup></td>
<td align="center" valign="top">7.0&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;9</sup> (96%)</td>
<td align="center" valign="top">1.3&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (100%)</td>
<td align="center" valign="top">C&#x02013;D</td>
</tr>
<tr>
<td align="left" valign="top">8D12</td>
<td align="center" valign="top">5.1&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;4</sup></td>
<td align="center" valign="top">6.1&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;11</sup></td>
<td align="center" valign="top">1.1&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (98%)</td>
<td align="center" valign="top">1.5&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (100%)</td>
<td align="center" valign="top">A</td>
</tr>
<tr>
<td align="left" valign="top">8C09</td>
<td align="center" valign="top">2.3&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;3</sup></td>
<td align="center" valign="top">1.7&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;7</sup></td>
<td align="center" valign="top">4.7&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (90%)</td>
<td align="center" valign="top">3.3&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (100%)</td>
<td align="center" valign="top">A</td>
</tr>
<tr>
<td align="left" valign="top">8G11</td>
<td align="center" valign="top">9.7&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;4</sup></td>
<td align="center" valign="top">3.4&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup></td>
<td align="center" valign="top">2.4&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;7</sup> (92%)</td>
<td align="center" valign="top">8.2&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (98%)</td>
<td align="center" valign="top">B</td>
</tr>
<tr>
<td align="left" valign="top">5C06</td>
<td align="center" valign="top">3.7&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;3</sup></td>
<td align="center" valign="top">2.0&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;7</sup></td>
<td align="center" valign="top">2.2&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (98%)</td>
<td align="center" valign="top">1.2&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;7</sup> (98%)</td>
<td align="center" valign="top">A</td>
</tr>
<tr>
<td align="left" valign="top">2C06</td>
<td align="center" valign="top">6.9&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;3</sup></td>
<td align="center" valign="top">&#x0003E;1.0&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;6</sup></td>
<td align="center" valign="top">2.5&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;7</sup> (90%)</td>
<td align="center" valign="top">3.0&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;7</sup> (92%)</td>
<td align="center" valign="top">A</td>
</tr>
<tr>
<td align="left" valign="top">5G04</td>
<td align="center" valign="top">2.9&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;4</sup></td>
<td align="center" valign="top">3.3&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;10</sup></td>
<td align="center" valign="top">n.a. (92%)</td>
<td align="center" valign="top">4.9&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;9</sup> (100%)</td>
<td align="center" valign="top">C&#x02013;D</td>
</tr>
<tr>
<td align="left" valign="top">2D07</td>
<td align="center" valign="top">5.0&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;3</sup></td>
<td align="center" valign="top">5.0&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;9</sup></td>
<td align="center" valign="top">n.a. (30%)</td>
<td align="center" valign="top">1.8&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (100%)</td>
<td align="center" valign="top">D</td>
</tr>
<tr>
<td align="left" valign="top">2B09</td>
<td align="center" valign="top">1.8&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;3</sup></td>
<td align="center" valign="top">1.3&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;9</sup></td>
<td align="center" valign="top">n.a. (67%)</td>
<td align="center" valign="top">9.6&#x02009;&#x000D7;&#x02009;10<sup>&#x02212;8</sup> (96%)</td>
<td align="center" valign="top">C</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn id="tfn1"><p><italic><sup>a</sup>Efficacy or maximum inhibition is shown between parentheses</italic>.</p></fn><p><italic>n.a.: IC<sub>50</sub> values could not be determined due to incomplete dose&#x02013;responses in the range of concentrations tested. The reported inhibition corresponds to the % inhibition observed at 1&#x02009;&#x000B5;M of 2D07 and 2&#x02009;&#x000B5;M of 5G04 and 2B09</italic>.</p></table-wrap-foot></table-wrap>
<fig id="F6" position="float">
<label>Figure 6</label>
<caption><p><bold>Dose&#x02013;response curves of selected anti-human RON (hRON) nanobodies inhibiting ligand-induced ERK phosphorylation (A) and binding of ligand to hRON (B)</bold>. Symbol colors relate to the different epitope bins to which the nanobodies belong (Table <xref ref-type="table" rid="T3">3</xref>): bin A (shades of blue), bin B (purple), bin C (green), bin D (red), and bin C&#x02013;D (shades of orange).</p></caption>
<graphic xlink:href="fimmu-08-00420-g006.tif"/>
</fig>
</sec>
<sec id="S3-4">
<title>Sequence Diversity of RON Nanobodies</title>
<p>Sequence analysis of the 28 randomly selected nanobodies and the 35 nanobodies identified by both conventional and NGS campaigns revealed an extensive functional sequence diversity (Figure <xref ref-type="fig" rid="F7">7</xref>). This is best illustrated by the observation that most nanobodies have very different CDR sequences. Together, these results confirm that our NGS-based approach is able to correctly predict large numbers of unrelated functional nanobody sequences targeting the same antigen and illustrate the functional diversity and quality of the outbred camelid&#x02019;s heavy chain-only immune response.</p>
<fig id="F7" position="float">
<label>Figure 7</label>
<caption><p><bold>Alignment of human RON (hRON) nanobodies (see also Table S2 in Supplementary Material)</bold>. The 28 randomly selected candidate hRON-binding nanobodies are identified by the acronym &#x0201C;NGS&#x0201D; followed by a five digit number. The three sequences marked by an asterisk (NGS00003, NGS00020, and NGS00027) are the non-binding sequences from the randomly selected panel of 28. The 35 nanobodies discovered in the conventional screening campaign and predicted by the next-generation sequencing (NGS) analysis are identified by a one or two digit number, followed by a letter, followed by a two digit number. Numbering of alignment positions was done according to the IMGT V-DOMAIN system (<xref ref-type="bibr" rid="B26">26</xref>). CDR regions are highlighted in gray. Dots represent residues identical to the top sequence. Dashes represent gaps introduced by the alignment.</p></caption>
<graphic xlink:href="fimmu-08-00420-g007.tif"/>
</fig>
</sec>
</sec>
<sec id="S4" sec-type="discussion">
<title>Discussion</title>
<p>A classic difficulty in the field of antibody repertoire sequencing is the clustering of clonally related sequences derived from the same progenitor during B cell maturation (<xref ref-type="bibr" rid="B12">12</xref>, <xref ref-type="bibr" rid="B13">13</xref>). While NGS analysis for antibody-derived binders such as scFvs and Fabs often is limited to the CDR3 region, the short length of nanobodies brings the advantage to obtain high quality full-length coverage by pairing of forward and reverse reads obtained with Illumina 2&#x02009;&#x000D7;&#x02009;250&#x02009;bp chemistry, as demonstrated before (<xref ref-type="bibr" rid="B17">17</xref>&#x02013;<xref ref-type="bibr" rid="B21">21</xref>). As a consequence, the downstream data analysis can reliably make use of all the FR and CDR sequence information. Without experimental data to support relatedness of antibody sequences at the phenotypic level, selecting a sequence identity threshold for clustering is relatively arbitrary. Here, we applied an inverse approach to the problem: rather than defining relatedness, we sought to define unrelatedness. Using a large set of publicly available nanobody sequences, we explored a range of sequence identity thresholds. The diverse nature of this data set makes it a highly representative source to sample unrelatedness. Clustering of unrelated nanobody sequences became apparent at sequence identity thresholds of 80% and lower. We selected a threshold of 90% to cluster the NGS dataset and subsequently obtained a high degree of experimental validation for these clusters. Other sequence identity thresholds could of course be explored, involving a tradeoff between the number of candidate clusters and their relative correctness. An increased stringency results in a larger number of clusters containing fewer unrelated sequences, whereas a lower stringency results in fewer clusters to choose from, with a higher proportion of unrelated sequences. The concept of using the diversity of publicly available data to establish meaningful sequence identity thresholds for clustering of related sequences is applicable to other types of antibody-derived binding domains and simple binding scaffolds.</p>
<p>The major challenge was the choice of inclusion criteria to representatively sample such a large diversity of candidate binders. One way to reduce the number of candidate binders is to apply more stringent cutoff values to cluster size and enrichment factor. However, this creates a bias toward the more abundant binders which are also identified by the conventional screening approach, as shown here. More so, we did not observe a clear correlation between the binding properties of a given cluster and its size or enrichment factor. In other words, good binders can be found among the more abundant and enriched clusters as well as among the less frequent clusters. Alternatively, lowering the sequence identity threshold for clustering would result in a lower number of clusters to sample from. However, as discussed above, this would increase the likelihood of clustering unrelated sequences, resulting in a higher proportion of erratic majority-rule consensuses as representatives. By random sampling representatively across a wide range of cluster sizes and enrichment factors, we achieved around 90% success rate in identifying anti-hRON nanobodies with binding characteristics and functional blockade comparable to those of the conventional screening campaign. Roughly half of these binders functionally inhibited hRON signaling. As such, it appears reasonable to assume that a large fraction of the &#x0003E;5,000 other enriched sequences qualify as nanobodies functionally blocking hRON.</p>
<p>The abovementioned high success rate also validates the combination of clustering related sequences and majority-rule consensus building as a very effective method to deal with PCR-induced and NGS read errors.</p>
<p>An alternative approach could be envisaged, leaving out the negative control sample, whereby sequencing and data analysis costs would be halved. When applying cluster size &#x0003E;10 as the inclusion criterion, this would increase the number of clusters-of-interest in the RON sample from 5,173 to 1.2&#x02009;&#x000D7;&#x02009;10<sup>4</sup> (Table <xref ref-type="table" rid="T2">2</xref>). Although a large fraction of these clusters can be expected to be enriched and functional, it is reasonable to assume that a fair number of these would be enriched by the phage display selections for the wrong reasons (display efficiency, stickiness, off-target binding). As a result, the fraction of false positives would be higher in comparison to an NGS approach including a proper negative control phage display sample. Hence, the upstream sequencing and analysis cost savings could be offset by an increase in downstream gene synthesis and screening costs.</p>
<p>Twelve of the hRON-binding nanobodies that were further characterized cover four different non-overlapping epitopes. Most of them inhibit MSP ligand binding to hRON and all fully block downstream ERK phosphorylation. This limited sample represents an interesting mix of ligand-dependent and -independent modes-of-action for the blocking of RON signaling. Many examples document the ability of the outbred camelid&#x02019;s immune response to generate nanobodies against challenging targets including ion channels (<xref ref-type="bibr" rid="B29">29</xref>&#x02013;<xref ref-type="bibr" rid="B31">31</xref>), GPCRs (<xref ref-type="bibr" rid="B32">32</xref>, <xref ref-type="bibr" rid="B33">33</xref>), small molecules and toxins (<xref ref-type="bibr" rid="B34">34</xref>, <xref ref-type="bibr" rid="B35">35</xref>), viruses (<xref ref-type="bibr" rid="B34">34</xref>, <xref ref-type="bibr" rid="B36">36</xref>) and enzymes (<xref ref-type="bibr" rid="B2">2</xref>). To our knowledge, this is the first example to illustrate the potential extent of an outbred camelid&#x02019;s functional immune response in terms of sequence diversity.</p>
<p>In conclusion, an NGS-based discovery approach combining full-length sequence clustering and the use of majority-rule consensuses as representatives reveals a highly diverse landscape of selective, functional nanobodies.</p>
</sec>
<sec id="S5">
<title>Ethics Statement</title>
<p>This study was carried out in accordance to EU animal welfare legislation and after approval of the local ethics committee &#x0201C;Ethical Committee Ablynx Camelid Facility LA1400575.&#x0201D;</p>
</sec>
<sec id="S6" sec-type="author-contributor">
<title>Author Contributions</title>
<p>DF, RM, JC, JA, YB, RF, DR, RT, DT, and LV performed experiments. AV, PD, MC, CS, and BD conceived and designed experiments and analyzed data. ML, PD, and BD designed and build the bioinformatics pipeline. PD, AV, CS, and BD wrote the manuscript. All authors read and critically reviewed the manuscript.</p>
</sec>
<sec id="S7">
<title>Conflict of Interest Statement</title>
<p>All authors are or have been employees of Ablynx N.V.</p>
</sec>
</body>
<back>
<ack>
<p>The authors want to thank Samuel Arvidsson and Berthold Fartmann at LGC Genomics for their handling of the Illumina MiSeq samples and raw data processing.</p>
</ack>
<sec id="S8">
<title>Funding</title>
<p>This work was supported by Ablynx N.V.</p>
</sec>
<sec id="S9" sec-type="supplementary-material">
<title>Supplementary Material</title>
<p>The Supplementary Material for this article can be found online at <uri xlink:href="http://journal.frontiersin.org/article/10.3389/fimmu.2017.00420/full&#x00023;supplementary-material">http://journal.frontiersin.org/article/10.3389/fimmu.2017.00420/full&#x00023;supplementary-material</uri>.</p>
<supplementary-material xlink:href="Table_1.XLSX" id="SM1" mimetype="applicationn/XLSX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table_2.DOCX" id="SM2" mimetype="applicationn/DOCX" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image_1.PDF" id="SM3" mimetype="applicationn/PDF" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1"><label>1</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Steeland</surname> <given-names>S</given-names></name> <name><surname>Vandenbroucke</surname> <given-names>RE</given-names></name> <name><surname>Libert</surname> <given-names>C</given-names></name></person-group>. <article-title>Nanobodies as therapeutics: big opportunities for small antibodies</article-title>. <source>Drug Discov Today</source> (<year>2016</year>) <volume>21</volume>:<fpage>1076</fpage>&#x02013;<lpage>113</lpage>.<pub-id pub-id-type="doi">10.1016/j.drudis.2016.04.003</pub-id><pub-id pub-id-type="pmid">27080147</pub-id></citation></ref>
<ref id="B2"><label>2</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Muyldermans</surname> <given-names>S</given-names></name></person-group>. <article-title>Nanobodies: natural single-domain antibodies</article-title>. <source>Annu Rev Biochem</source> (<year>2013</year>) <volume>82</volume>:<fpage>775</fpage>&#x02013;<lpage>97</lpage>.<pub-id pub-id-type="doi">10.1146/annurev-biochem-063011-092449</pub-id><pub-id pub-id-type="pmid">23495938</pub-id></citation></ref>
<ref id="B3"><label>3</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pardon</surname> <given-names>E</given-names></name> <name><surname>Laeremans</surname> <given-names>T</given-names></name> <name><surname>Triest</surname> <given-names>S</given-names></name> <name><surname>Rasmussen</surname> <given-names>SGF</given-names></name> <name><surname>Wohlk&#x000F6;nig</surname> <given-names>A</given-names></name> <name><surname>Ruf</surname> <given-names>A</given-names></name> <etal/></person-group> <article-title>A general protocol for the generation of nanobodies for structural biology</article-title>. <source>Nat Protoc</source> (<year>2014</year>) <volume>9</volume>:<fpage>674</fpage>&#x02013;<lpage>93</lpage>.<pub-id pub-id-type="doi">10.1038/nprot.2014.039</pub-id><pub-id pub-id-type="pmid">24577359</pub-id></citation></ref>
<ref id="B4"><label>4</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bencurova</surname> <given-names>E</given-names></name> <name><surname>Pulzova</surname> <given-names>L</given-names></name> <name><surname>Flachbartova</surname> <given-names>Z</given-names></name> <name><surname>Bhide</surname> <given-names>M</given-names></name></person-group>. <article-title>A rapid and simple pipeline for synthesis of mRNA-ribosome-V(H)H complexes used in single-domain antibody ribosome display</article-title>. <source>Mol Biosyst</source> (<year>2015</year>) <volume>11</volume>:<fpage>1515</fpage>&#x02013;<lpage>24</lpage>.<pub-id pub-id-type="doi">10.1039/c5mb00026b</pub-id><pub-id pub-id-type="pmid">25902394</pub-id></citation></ref>
<ref id="B5"><label>5</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fleetwood</surname> <given-names>F</given-names></name> <name><surname>Devoogdt</surname> <given-names>N</given-names></name> <name><surname>Pellis</surname> <given-names>M</given-names></name> <name><surname>Wernery</surname> <given-names>U</given-names></name> <name><surname>Muyldermans</surname> <given-names>S</given-names></name> <name><surname>St&#x000E5;hl</surname> <given-names>S</given-names></name> <etal/></person-group> <article-title>Surface display of a single-domain antibody library on Gram-positive bacteria</article-title>. <source>Cell Mol Life Sci</source> (<year>2013</year>) <volume>70</volume>:<fpage>1081</fpage>&#x02013;<lpage>93</lpage>.<pub-id pub-id-type="doi">10.1007/s00018-012-1179-y</pub-id><pub-id pub-id-type="pmid">23064703</pub-id></citation></ref>
<ref id="B6"><label>6</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Koide</surname> <given-names>A</given-names></name> <name><surname>Koide</surname> <given-names>S</given-names></name></person-group>. <article-title>Affinity maturation of single-domain antibodies by yeast surface display</article-title>. <source>Methods Mol Biol</source> (<year>2012</year>) <volume>911</volume>:<fpage>431</fpage>&#x02013;<lpage>43</lpage>.<pub-id pub-id-type="doi">10.1007/978-1-61779-968-6_26</pub-id></citation></ref>
<ref id="B7"><label>7</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>De Schutter</surname> <given-names>K</given-names></name> <name><surname>Callewaert</surname> <given-names>N</given-names></name></person-group>. <article-title>Pichia surface display: a tool for screening single domain antibodies</article-title>. <source>Methods Mol Biol</source> (<year>2012</year>) <volume>911</volume>:<fpage>125</fpage>&#x02013;<lpage>34</lpage>.<pub-id pub-id-type="doi">10.1007/978-1-61779-968-6_8</pub-id><pub-id pub-id-type="pmid">22886249</pub-id></citation></ref>
<ref id="B8"><label>8</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ryckaert</surname> <given-names>S</given-names></name> <name><surname>Pardon</surname> <given-names>E</given-names></name> <name><surname>Steyaert</surname> <given-names>J</given-names></name> <name><surname>Callewaert</surname> <given-names>N</given-names></name></person-group>. <article-title>Isolation of antigen-binding camelid heavy chain antibody fragments (nanobodies) from an immune library displayed on the surface of <italic>Pichia pastoris</italic></article-title>. <source>J Biotechnol</source> (<year>2010</year>) <volume>145</volume>:<fpage>93</fpage>&#x02013;<lpage>8</lpage>.<pub-id pub-id-type="doi">10.1016/j.jbiotec.2009.10.010</pub-id><pub-id pub-id-type="pmid">19861136</pub-id></citation></ref>
<ref id="B9"><label>9</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pellis</surname> <given-names>M</given-names></name> <name><surname>Muyldermans</surname> <given-names>S</given-names></name> <name><surname>Vincke</surname> <given-names>C</given-names></name></person-group>. <article-title>Bacterial two hybrid: a versatile one-step intracellular selection method</article-title>. <source>Methods Mol Biol</source> (<year>2012</year>) <volume>911</volume>:<fpage>135</fpage>&#x02013;<lpage>50</lpage>.<pub-id pub-id-type="doi">10.1007/978-1-61779-968-6_9</pub-id><pub-id pub-id-type="pmid">22886250</pub-id></citation></ref>
<ref id="B10"><label>10</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gao</surname> <given-names>X</given-names></name> <name><surname>Hu</surname> <given-names>X</given-names></name> <name><surname>Tong</surname> <given-names>L</given-names></name> <name><surname>Liu</surname> <given-names>D</given-names></name> <name><surname>Chang</surname> <given-names>X</given-names></name> <name><surname>Wang</surname> <given-names>H</given-names></name> <etal/></person-group> <article-title>Construction of a camelid VHH yeast two-hybrid library and the selection of VHH against haemagglutinin-neuraminidase protein of the Newcastle disease virus</article-title>. <source>BMC Vet Res</source> (<year>2016</year>) <volume>12</volume>:<fpage>39</fpage>.<pub-id pub-id-type="doi">10.1186/s12917-016-0664-1</pub-id><pub-id pub-id-type="pmid">26920806</pub-id></citation></ref>
<ref id="B11"><label>11</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Glanville</surname> <given-names>J</given-names></name> <name><surname>D&#x02019;Angelo</surname> <given-names>S</given-names></name> <name><surname>Khan</surname> <given-names>TA</given-names></name> <name><surname>Reddy</surname> <given-names>ST</given-names></name> <name><surname>Naranjo</surname> <given-names>L</given-names></name> <name><surname>Ferrara</surname> <given-names>F</given-names></name> <etal/></person-group> <article-title>Deep sequencing in library selection projects: what insight does it bring?</article-title> <source>Curr Opin Struct Biol</source> (<year>2015</year>) <volume>33</volume>:<fpage>146</fpage>&#x02013;<lpage>60</lpage>.<pub-id pub-id-type="doi">10.1016/j.sbi.2015.09.001</pub-id><pub-id pub-id-type="pmid">26451649</pub-id></citation></ref>
<ref id="B12"><label>12</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hoehn</surname> <given-names>KB</given-names></name> <name><surname>Fowler</surname> <given-names>A</given-names></name> <name><surname>Lunter</surname> <given-names>G</given-names></name> <name><surname>Pybus</surname> <given-names>OG</given-names></name></person-group>. <article-title>The diversity and molecular evolution of B-cell receptors during infection</article-title>. <source>Mol Biol Evol</source> (<year>2016</year>) <volume>33</volume>:<fpage>1147</fpage>&#x02013;<lpage>57</lpage>.<pub-id pub-id-type="doi">10.1093/molbev/msw015</pub-id></citation></ref>
<ref id="B13"><label>13</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yaari</surname> <given-names>G</given-names></name> <name><surname>Kleinstein</surname> <given-names>SH</given-names></name></person-group>. <article-title>Practical guidelines for B-cell receptor repertoire sequencing analysis</article-title>. <source>Genome Med</source> (<year>2015</year>) <volume>7</volume>:<fpage>121</fpage>.<pub-id pub-id-type="doi">10.1186/s13073-015-0243-2</pub-id><pub-id pub-id-type="pmid">26589402</pub-id></citation></ref>
<ref id="B14"><label>14</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ravn</surname> <given-names>U</given-names></name> <name><surname>Gueneau</surname> <given-names>F</given-names></name> <name><surname>Baerlocher</surname> <given-names>L</given-names></name> <name><surname>Osteras</surname> <given-names>M</given-names></name> <name><surname>Desmurs</surname> <given-names>M</given-names></name> <name><surname>Malinge</surname> <given-names>P</given-names></name> <etal/></person-group> <article-title>By-passing in vitro screening &#x02013; next generation sequencing technologies applied to antibody display and in silico candidate selection</article-title>. <source>Nucleic Acids Res</source> (<year>2010</year>) <volume>38</volume>:<fpage>e193</fpage>.<pub-id pub-id-type="doi">10.1093/nar/gkq789</pub-id></citation></ref>
<ref id="B15"><label>15</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ravn</surname> <given-names>U</given-names></name> <name><surname>Didelot</surname> <given-names>G</given-names></name> <name><surname>Venet</surname> <given-names>S</given-names></name> <name><surname>Ng</surname> <given-names>K-T</given-names></name> <name><surname>Gueneau</surname> <given-names>F</given-names></name> <name><surname>Rousseau</surname> <given-names>F</given-names></name> <etal/></person-group> <article-title>Deep sequencing of phage display libraries to support antibody discovery</article-title>. <source>Methods</source> (<year>2013</year>) <volume>60</volume>:<fpage>99</fpage>&#x02013;<lpage>110</lpage>.<pub-id pub-id-type="doi">10.1016/j.ymeth.2013.03.001</pub-id></citation></ref>
<ref id="B16"><label>16</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Naso</surname> <given-names>MF</given-names></name> <name><surname>Lu</surname> <given-names>J</given-names></name> <name><surname>Panavas</surname> <given-names>T</given-names></name></person-group>. <article-title>Deep sequencing approaches to antibody discovery</article-title>. <source>Curr Drug Discov Technol</source> (<year>2014</year>) <volume>11</volume>:<fpage>85</fpage>&#x02013;<lpage>95</lpage>.<pub-id pub-id-type="doi">10.2174/15701638113106660040</pub-id><pub-id pub-id-type="pmid">24020911</pub-id></citation></ref>
<ref id="B17"><label>17</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fridy</surname> <given-names>PC</given-names></name> <name><surname>Li</surname> <given-names>Y</given-names></name> <name><surname>Keegan</surname> <given-names>S</given-names></name> <name><surname>Thompson</surname> <given-names>MK</given-names></name> <name><surname>Nudelman</surname> <given-names>I</given-names></name> <name><surname>Scheid</surname> <given-names>JF</given-names></name> <etal/></person-group> <article-title>A robust pipeline for rapid production of versatile nanobody repertoires</article-title>. <source>Nat Methods</source> (<year>2014</year>) <volume>11</volume>:<fpage>1253</fpage>&#x02013;<lpage>60</lpage>.<pub-id pub-id-type="doi">10.1038/nmeth.3170</pub-id><pub-id pub-id-type="pmid">25362362</pub-id></citation></ref>
<ref id="B18"><label>18</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Henry</surname> <given-names>KA</given-names></name> <name><surname>Tanha</surname> <given-names>J</given-names></name> <name><surname>Hussack</surname> <given-names>G</given-names></name></person-group>. <article-title>Identification of cross-reactive single-domain antibodies against serum albumin using next-generation DNA sequencing</article-title>. <source>Protein Eng Des Sel</source> (<year>2015</year>) <volume>28</volume>:<fpage>379</fpage>&#x02013;<lpage>83</lpage>.<pub-id pub-id-type="doi">10.1093/protein/gzv039</pub-id><pub-id pub-id-type="pmid">26319004</pub-id></citation></ref>
<ref id="B19"><label>19</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Miyazaki</surname> <given-names>N</given-names></name> <name><surname>Kiyose</surname> <given-names>N</given-names></name> <name><surname>Akazawa</surname> <given-names>Y</given-names></name> <name><surname>Takashima</surname> <given-names>M</given-names></name> <name><surname>Hagihara</surname> <given-names>Y</given-names></name> <name><surname>Inoue</surname> <given-names>N</given-names></name> <etal/></person-group> <article-title>Isolation and characterization of antigen-specific alpaca (<italic>Lama pacos</italic>) VHH antibodies by biopanning followed by high-throughput sequencing</article-title>. <source>J Biochem</source> (<year>2015</year>) <volume>158</volume>:<fpage>205</fpage>&#x02013;<lpage>15</lpage>.<pub-id pub-id-type="doi">10.1093/jb/mvv038</pub-id><pub-id pub-id-type="pmid">25888581</pub-id></citation></ref>
<ref id="B20"><label>20</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Henry</surname> <given-names>KA</given-names></name> <name><surname>Hussack</surname> <given-names>G</given-names></name> <name><surname>Collins</surname> <given-names>C</given-names></name> <name><surname>Zwaagstra</surname> <given-names>JC</given-names></name> <name><surname>Tanha</surname> <given-names>J</given-names></name> <name><surname>MacKenzie</surname> <given-names>CR</given-names></name></person-group>. <article-title>Isolation of TGF-&#x003B2;-neutralizing single-domain antibodies of predetermined epitope specificity using next-generation DNA sequencing</article-title>. <source>Protein Eng Des Sel</source> (<year>2016</year>) <volume>29</volume>:<fpage>439</fpage>&#x02013;<lpage>43</lpage>.<pub-id pub-id-type="doi">10.1093/protein/gzw043</pub-id><pub-id pub-id-type="pmid">27613412</pub-id></citation></ref>
<ref id="B21"><label>21</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Turner</surname> <given-names>KB</given-names></name> <name><surname>Naciri</surname> <given-names>J</given-names></name> <name><surname>Liu</surname> <given-names>JL</given-names></name> <name><surname>Anderson</surname> <given-names>GP</given-names></name> <name><surname>Goldman</surname> <given-names>ER</given-names></name> <name><surname>Zabetakis</surname> <given-names>D</given-names></name></person-group>. <article-title>Next-generation sequencing of a single domain antibody repertoire reveals quality of phage display selected candidates</article-title>. <source>PLoS One</source> (<year>2016</year>) <volume>11</volume>:<fpage>e0149393</fpage>.<pub-id pub-id-type="doi">10.1371/journal.pone.0149393</pub-id><pub-id pub-id-type="pmid">26895405</pub-id></citation></ref>
<ref id="B22"><label>22</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ronsin</surname> <given-names>C</given-names></name> <name><surname>Muscatelli</surname> <given-names>F</given-names></name> <name><surname>Mattei</surname> <given-names>MG</given-names></name> <name><surname>Breathnach</surname> <given-names>R</given-names></name></person-group>. <article-title>A novel putative receptor protein tyrosine kinase of the met family</article-title>. <source>Oncogene</source> (<year>1993</year>) <volume>8</volume>:<fpage>1195</fpage>&#x02013;<lpage>202</lpage>.<pub-id pub-id-type="pmid">8386824</pub-id></citation></ref>
<ref id="B23"><label>23</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gaudino</surname> <given-names>G</given-names></name> <name><surname>Follenzi</surname> <given-names>A</given-names></name> <name><surname>Naldini</surname> <given-names>L</given-names></name> <name><surname>Collesi</surname> <given-names>C</given-names></name> <name><surname>Santoro</surname> <given-names>M</given-names></name> <name><surname>Gallo</surname> <given-names>KA</given-names></name> <etal/></person-group> <article-title>RON is a heterodimeric tyrosine kinase receptor activated by the HGF homologue MSP</article-title>. <source>EMBO J</source> (<year>1994</year>) <volume>13</volume>:<fpage>3524</fpage>&#x02013;<lpage>32</lpage>.<pub-id pub-id-type="pmid">8062829</pub-id></citation></ref>
<ref id="B24"><label>24</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>M-H</given-names></name> <name><surname>Zhang</surname> <given-names>R</given-names></name> <name><surname>Zhou</surname> <given-names>Y-Q</given-names></name> <name><surname>Yao</surname> <given-names>H-P</given-names></name></person-group>. <article-title>Pathogenesis of RON receptor tyrosine kinase in cancer cells: activation mechanism, functional crosstalk, and signaling addiction</article-title>. <source>J Biomed Res</source> (<year>2013</year>) <volume>27</volume>:<fpage>345</fpage>&#x02013;<lpage>56</lpage>.<pub-id pub-id-type="doi">10.7555/JBR.27.20130038</pub-id><pub-id pub-id-type="pmid">24086167</pub-id></citation></ref>
<ref id="B25"><label>25</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mago&#x0010D;</surname> <given-names>T</given-names></name> <name><surname>Salzberg</surname> <given-names>SL</given-names></name></person-group>. <article-title>FLASH: fast length adjustment of short reads to improve genome assemblies</article-title>. <source>Bioinformatics</source> (<year>2011</year>) <volume>27</volume>:<fpage>2957</fpage>&#x02013;<lpage>63</lpage>.<pub-id pub-id-type="doi">10.1093/bioinformatics/btr507</pub-id><pub-id pub-id-type="pmid">21903629</pub-id></citation></ref>
<ref id="B26"><label>26</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lefranc</surname> <given-names>M-P</given-names></name> <name><surname>Pommi&#x000E9;</surname> <given-names>C</given-names></name> <name><surname>Ruiz</surname> <given-names>M</given-names></name> <name><surname>Giudicelli</surname> <given-names>V</given-names></name> <name><surname>Foulquier</surname> <given-names>E</given-names></name> <name><surname>Truong</surname> <given-names>L</given-names></name> <etal/></person-group> <article-title>IMGT unique numbering for immunoglobulin and T cell receptor variable domains and Ig superfamily V-like domains</article-title>. <source>Dev Comp Immunol</source> (<year>2003</year>) <volume>27</volume>:<fpage>55</fpage>&#x02013;<lpage>77</lpage>.<pub-id pub-id-type="doi">10.1016/S0145-305X(02)00039-3</pub-id><pub-id pub-id-type="pmid">12477501</pub-id></citation></ref>
<ref id="B27"><label>27</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>W</given-names></name> <name><surname>Godzik</surname> <given-names>A</given-names></name></person-group>. <article-title>Cd-hit: a fast program for clustering and comparing large sets of protein or nucleotide sequences</article-title>. <source>Bioinformatics</source> (<year>2006</year>) <volume>22</volume>:<fpage>1658</fpage>&#x02013;<lpage>9</lpage>.<pub-id pub-id-type="doi">10.1093/bioinformatics/btl158</pub-id><pub-id pub-id-type="pmid">16731699</pub-id></citation></ref>
<ref id="B28"><label>28</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fu</surname> <given-names>L</given-names></name> <name><surname>Niu</surname> <given-names>B</given-names></name> <name><surname>Zhu</surname> <given-names>Z</given-names></name> <name><surname>Wu</surname> <given-names>S</given-names></name> <name><surname>Li</surname> <given-names>W</given-names></name></person-group>. <article-title>CD-HIT: accelerated for clustering the next-generation sequencing data</article-title>. <source>Bioinformatics</source> (<year>2012</year>) <volume>28</volume>:<fpage>3150</fpage>&#x02013;<lpage>2</lpage>.<pub-id pub-id-type="doi">10.1093/bioinformatics/bts565</pub-id><pub-id pub-id-type="pmid">23060610</pub-id></citation></ref>
<ref id="B29"><label>29</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wei</surname> <given-names>G</given-names></name> <name><surname>Meng</surname> <given-names>W</given-names></name> <name><surname>Guo</surname> <given-names>H</given-names></name> <name><surname>Pan</surname> <given-names>W</given-names></name> <name><surname>Liu</surname> <given-names>J</given-names></name> <name><surname>Peng</surname> <given-names>T</given-names></name> <etal/></person-group> <article-title>Potent neutralization of influenza A virus by a single-domain antibody blocking M2 ion channel protein</article-title>. <source>PLoS One</source> (<year>2011</year>) <volume>6</volume>:<fpage>e28309</fpage>.<pub-id pub-id-type="doi">10.1371/journal.pone.0028309</pub-id><pub-id pub-id-type="pmid">22164266</pub-id></citation></ref>
<ref id="B30"><label>30</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hassiki</surname> <given-names>R</given-names></name> <name><surname>Labro</surname> <given-names>AJ</given-names></name> <name><surname>Benlasfar</surname> <given-names>Z</given-names></name> <name><surname>Vincke</surname> <given-names>C</given-names></name> <name><surname>Somia</surname> <given-names>M</given-names></name> <name><surname>El Ayeb</surname> <given-names>M</given-names></name> <etal/></person-group> <article-title>Dromedary immune response and specific Kv2.1 antibody generation using a specific immunization approach</article-title>. <source>Int J Biol Macromol</source> (<year>2016</year>) <volume>93</volume>:<fpage>167</fpage>&#x02013;<lpage>71</lpage>.<pub-id pub-id-type="doi">10.1016/j.ijbiomac.2016.06.031</pub-id><pub-id pub-id-type="pmid">27320844</pub-id></citation></ref>
<ref id="B31"><label>31</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Danquah</surname> <given-names>W</given-names></name> <name><surname>Meyer-Schwesinger</surname> <given-names>C</given-names></name> <name><surname>Rissiek</surname> <given-names>B</given-names></name> <name><surname>Pinto</surname> <given-names>C</given-names></name> <name><surname>Serracant-Prat</surname> <given-names>A</given-names></name> <name><surname>Amadi</surname> <given-names>M</given-names></name> <etal/></person-group> <article-title>Nanobodies that block gating of the P2X7 ion channel ameliorate inflammation</article-title>. <source>Sci Transl Med</source> (<year>2016</year>) <volume>8</volume>:<fpage>366ra162</fpage>.<pub-id pub-id-type="doi">10.1126/scitranslmed.aaf8463</pub-id></citation></ref>
<ref id="B32"><label>32</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Muji&#x00107;-Deli&#x00107;</surname> <given-names>A</given-names></name> <name><surname>de Wit</surname> <given-names>RH</given-names></name> <name><surname>Verkaar</surname> <given-names>F</given-names></name> <name><surname>Smit</surname> <given-names>MJ</given-names></name></person-group>. <article-title>GPCR-targeting nanobodies: attractive research tools, diagnostics, and therapeutics</article-title>. <source>Trends Pharmacol Sci</source> (<year>2014</year>) <volume>35</volume>:<fpage>247</fpage>&#x02013;<lpage>55</lpage>.<pub-id pub-id-type="doi">10.1016/j.tips.2014.03.003</pub-id><pub-id pub-id-type="pmid">24690241</pub-id></citation></ref>
<ref id="B33"><label>33</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cromie</surname> <given-names>KD</given-names></name> <name><surname>Van Heeke</surname> <given-names>G</given-names></name> <name><surname>Boutton</surname> <given-names>C</given-names></name></person-group>. <article-title>Nanobodies and their use in GPCR drug discovery</article-title>. <source>Curr Top Med Chem</source> (<year>2015</year>) <volume>15</volume>:<fpage>2543</fpage>&#x02013;<lpage>57</lpage>.<pub-id pub-id-type="doi">10.2174/1568026615666150701113549</pub-id><pub-id pub-id-type="pmid">26126902</pub-id></citation></ref>
<ref id="B34"><label>34</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wesolowski</surname> <given-names>J</given-names></name> <name><surname>Alzogaray</surname> <given-names>V</given-names></name> <name><surname>Reyelt</surname> <given-names>J</given-names></name> <name><surname>Unger</surname> <given-names>M</given-names></name> <name><surname>Juarez</surname> <given-names>K</given-names></name> <name><surname>Urrutia</surname> <given-names>M</given-names></name> <etal/></person-group> <article-title>Single domain antibodies: promising experimental and therapeutic tools in infection and immunity</article-title>. <source>Med Microbiol Immunol</source> (<year>2009</year>) <volume>198</volume>:<fpage>157</fpage>&#x02013;<lpage>74</lpage>.<pub-id pub-id-type="doi">10.1007/s00430-009-0116-7</pub-id><pub-id pub-id-type="pmid">19529959</pub-id></citation></ref>
<ref id="B35"><label>35</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bever</surname> <given-names>CS</given-names></name> <name><surname>Dong</surname> <given-names>J-X</given-names></name> <name><surname>Vasylieva</surname> <given-names>N</given-names></name> <name><surname>Barnych</surname> <given-names>B</given-names></name> <name><surname>Cui</surname> <given-names>Y</given-names></name> <name><surname>Xu</surname> <given-names>Z-L</given-names></name> <etal/></person-group> <article-title>VHH antibodies: emerging reagents for the analysis of environmental chemicals</article-title>. <source>Anal Bioanal Chem</source> (<year>2016</year>) <volume>408</volume>:<fpage>5985</fpage>&#x02013;<lpage>6002</lpage>.<pub-id pub-id-type="doi">10.1007/s00216-016-9585-x</pub-id><pub-id pub-id-type="pmid">27209591</pub-id></citation></ref>
<ref id="B36"><label>36</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Vanlandschoot</surname> <given-names>P</given-names></name> <name><surname>Stortelers</surname> <given-names>C</given-names></name> <name><surname>Beirnaert</surname> <given-names>E</given-names></name> <name><surname>Iba&#x000F1;ez</surname> <given-names>LI</given-names></name> <name><surname>Schepens</surname> <given-names>B</given-names></name> <name><surname>Depla</surname> <given-names>E</given-names></name> <etal/></person-group> <article-title>Nanobodies<sup>&#x000AE;</sup>: new ammunition to battle viruses</article-title>. <source>Antiviral Res</source> (<year>2011</year>) <volume>92</volume>:<fpage>389</fpage>&#x02013;<lpage>407</lpage>.<pub-id pub-id-type="doi">10.1016/j.antiviral.2011.09.002</pub-id><pub-id pub-id-type="pmid">21939690</pub-id></citation></ref>
</ref-list>
<fn-group>
<fn id="fn1"><p><sup>1</sup><uri xlink:href="http://biopython.org/">http://biopython.org/</uri>.</p></fn>
<fn id="fn2"><p><sup>2</sup><uri xlink:href="http://www.ncbi.nlm.nih.gov/protein/">http://www.ncbi.nlm.nih.gov/protein/</uri>.</p></fn>
</fn-group>
</back>
</article>