<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Soft Matter</journal-id>
<journal-title>Frontiers in Soft Matter</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Soft Matter</abbrev-journal-title>
<issn pub-type="epub">2813-0499</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1402702</article-id>
<article-id pub-id-type="doi">10.3389/frsfm.2024.1402702</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Soft Matter</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Using QSAR to predict polymer-drug interactions for drug delivery</article-title>
<alt-title alt-title-type="left-running-head">Xin et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/frsfm.2024.1402702">10.3389/frsfm.2024.1402702</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Xin</surname>
<given-names>Alison W.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="fn" rid="fn1">
<sup>&#x2020;</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Rivera-Delgado</surname>
<given-names>Edgardo</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="fn" rid="fn1">
<sup>&#x2020;</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>von Recum</surname>
<given-names>Horst A.</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2578383/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Hathaway Brown High School</institution>, <institution>Case Western Reserve University</institution>, <addr-line>Cleveland</addr-line>, <addr-line>OH</addr-line>, <country>United States</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Department of Biomedical Engineering</institution>, <institution>Case Western Reserve University</institution>, <addr-line>Cleveland</addr-line>, <addr-line>OH</addr-line>, <country>United States</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/233004/overview">Ali Miserez</ext-link>, Nanyang Technological University, Singapore</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1403329/overview">Animesh Pan</ext-link>, University of Rhode Island, United States</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/970330/overview">Frank Alexis</ext-link>, Universidad San Francisco de Quito, Ecuador</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Horst A. von Recum, <email>horst.vonrecum@case.edu</email>
</corresp>
<fn fn-type="equal" id="fn1">
<label>
<sup>&#x2020;</sup>
</label>
<p>These authors have contributed equally to this work and share first authorship</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>29</day>
<month>07</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>4</volume>
<elocation-id>1402702</elocation-id>
<history>
<date date-type="received">
<day>18</day>
<month>03</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>20</day>
<month>06</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2024 Xin, Rivera-Delgado and von Recum.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Xin, Rivera-Delgado and von Recum</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Affinity-mediated drug delivery utilizes electrostatic, hydrophobic, or other non-covalent interactions between molecules and a polymer to extend the timeframe of drug release. Cyclodextrin polymers exhibit affinity interaction, however, experimentally testing drug candidates for affinity is time-consuming, making computational predictions more effective. One option, docking programs, provide predictions of affinity, but lack reliability, as their accuracy with cyclodextrin remains unverified experimentally. Alternatively, quantitative structure-activity relationship models (QSARs), which analyze statistical relationships between molecular properties, appear more promising. Previously constructed QSARs for cyclodextrin are not publicly available, necessitating an openly accessible model. Around 600 experimental affinities between cyclodextrin and guest molecules were cleaned and imported from published research. The software PaDEL-Descriptor calculated over 1,000 chemical descriptors for each molecule, which were then analyzed with R to create several QSARs with different statistical methods. These QSARs proved highly time efficient, calculating in minutes what docking programs could accomplish in hours. Additionally, on test sets, QSARs reached <italic>R</italic>
<sup>2</sup> values of around 0.7&#x2013;0.8. The speed, accuracy, and accessibility of these QSARs improve evaluation of individual drugs and facilitate screening of large datasets for potential candidates in cyclodextrin affinity-based delivery systems. An app was built to rapidly access model predictions for end users using the Shiny library. To demonstrate the usability for drug release planning, the QSAR predictions were coupled with a mechanistic model of diffusion within the app. Integrating new modules should provide an accessible approach to use other cheminformatic tools in the field of drug delivery.</p>
</abstract>
<kwd-group>
<kwd>QSPR (quantitative structure properties relationship)</kwd>
<kwd>drug delivery</kwd>
<kwd>cyclodextrin</kwd>
<kwd>machine learning (ML)</kwd>
<kwd>small molecules</kwd>
<kwd>ODE (ordinary differential equation)</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Polymers</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>Affinity delivery, which relies on interactions between a drug delivery system and drug molecules, improves effectiveness of medication by extending the duration of drug release and thereby lengthening the duration of the treatment (<xref ref-type="bibr" rid="B31">Rivera-Delgado et al., 2016</xref>). Mathematical modeling of these affinity systems has shown that the strength of the affinity interaction, the ratio of host binding sites to guest ligands, and the molecular path length of diffusion influence the transport of molecules out of the system. Of these physical forces, the affinity strength plays an important role in the classification of the system and the timescale of drug release (<xref ref-type="bibr" rid="B13">Fu et al., 2011</xref>). Affinity interaction can be associated with a variety of physical properties, including charge, hydrophobicity, Van der Waals forces, etc. In the fields of biomaterials and drug delivery, affinity delivery has been used with small molecule drugs (<xref ref-type="bibr" rid="B40">Wang and von Recum, 2011</xref>), proteins (<xref ref-type="bibr" rid="B31">Rivera-Delgado et al., 2016</xref>), cytokines, and antibodies (<xref ref-type="bibr" rid="B26">Ortiz et al., 2011</xref>).</p>
<p>Our lab tests rings of glucose molecules as affinity hosts called cyclodextrins, which are particularly promising affinity drug delivery hosts due to their structural properties, biocompatibility and versatility. The most common cyclodextrin are composed of a ring 6, 7, or 8 glucose molecules (&#x3b1;, &#x3b2;, and &#x3b3;-cyclodextrin, respectively), and the conformation of the hydroxyl groups of the ring create a basket-like structure with a hydrophobic interior and hydrophilic interior, allowing for complexation with drug molecules (<xref ref-type="fig" rid="F1">Figure 1</xref>). Additionally, cyclodextrin can be polymerized into a variety of materials, including microparticles, viscous gels, and solid films. Unfortunately, experiments to confirm sustained release from the affinity guest-host system often takes weeks, making testing large numbers of potential candidates for cyclodextrin release systems impractical.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Cyclodextrin complexation.</p>
</caption>
<graphic xlink:href="frsfm-04-1402702-g001.tif"/>
</fig>
<p>As an alternative to experimental testing, candidate molecules can be analyzed computationally. Predicting the binding affinity between cyclodextrin and drug molecules allows for the processing of molecules on the scale of minutes rather than weeks. There are two major methods for predicting molecular interaction: docking models and QSARs. Docking models use molecular force fields, which simulate interactions and potential energy between atoms. Force field parameters may be derived from experiments, calculations from quantum mechanics, or both (<xref ref-type="bibr" rid="B16">Jacob et al., 2012</xref>). In addition to providing a numeric estimate for binding affinity, docking programs produce visualizations of how molecules interact. QSARs, or Quantitative Structure-Activity Relationship models, statistically predict molecular interactions using molecular descriptors. Molecular descriptors are certain physical or chemical characteristics of molecules that can be evaluated numerically (for example, the number of hydrogen atoms or the length of the longest bond chain). Many different types of regression models and statistical learning methods can be used as QSARs, ranging in complexity from linear models to artificial neural networks (<xref ref-type="bibr" rid="B7">Dehmer et al., 2012</xref>).</p>
<p>Previous investigations have been made on the accuracy of both docking and QSARs in predicting cyclodextrin affinity, but examining a sample of these papers reveals several concerns (<xref ref-type="table" rid="T1">Table 1</xref>). Notably, all of the investigated models used software hidden behind a paywall or only available with a license (<xref ref-type="bibr" rid="B27">P&#xe9;rez-Garrido et al., 2009</xref>; <xref ref-type="bibr" rid="B28">Prakasvudhisarn et al., 2009</xref>; <xref ref-type="bibr" rid="B14">Ghasemi et al., 2011</xref>; <xref ref-type="bibr" rid="B21">Merzlikine et al., 2011</xref>; <xref ref-type="bibr" rid="B1">Ahmadi and Ghasemi, 2014</xref>; <xref ref-type="bibr" rid="B38">Veselinovi&#x107; et al., 2015</xref>; <xref ref-type="bibr" rid="B42">Xu et al., 2015</xref>; <xref ref-type="bibr" rid="B25">Mirrahimi et al., 2016</xref>). Additionally, many models lacked proper verification. Following Tropsha&#x2019;s publication detailing best practices for QSAR development, a completely verified model should undergo leave-one-out cross-validation (LOO-CV) (reported as Q<sup>2</sup>), y-randomization, pass a variety of internal accuracy tests, and be analyzed for applicability domain. Additionally, models should be evaluated on multiple test sets as well as a hold-out external validation set (<xref ref-type="bibr" rid="B36">Tropsha, 2010</xref>). Of the papers investigated, none contained the full set of verification strategies.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Results of previous cyclodextrin QSARs.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="2" align="left">QSAR</th>
<th rowspan="2" align="left">R2</th>
<th rowspan="2" align="left">Descriptors</th>
<th rowspan="2" align="left">Feature selection</th>
<th colspan="4" align="left">Validation</th>
</tr>
<tr>
<th align="left">
<italic>Q2</italic>
</th>
<th align="left">
<italic>y-rand</italic>
</th>
<th align="left">
<italic>AD</italic>
</th>
<th align="left">
<italic>EV</italic>
</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Cubist (<xref ref-type="bibr" rid="B14">Ghasemi et al., 2011</xref>)</td>
<td align="right">0.945</td>
<td align="left">Pfizer&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">Yes</td>
<td align="left">Yes</td>
</tr>
<tr>
<td align="left">Random forest (<xref ref-type="bibr" rid="B14">Ghasemi et al., 2011</xref>)</td>
<td align="right">0.912</td>
<td align="left">Pfizer&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">Yes</td>
<td align="left">Yes</td>
</tr>
<tr>
<td align="left">Partial least squares (PLS) (<xref ref-type="bibr" rid="B38">Veselinovi&#x107; et al., 2015</xref>)</td>
<td align="right">0.68</td>
<td align="left">ChemOffice, SYBYL, Pentacle&#x2a;</td>
<td align="left">Genetic algorithm</td>
<td align="right">0.64</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">Yes</td>
<td align="left">&#x2a;&#x2a;</td>
</tr>
<tr>
<td align="left">PLS (<xref ref-type="bibr" rid="B27">P&#xe9;rez-Garrido et al., 2009</xref>)</td>
<td align="right">0.74</td>
<td align="left">SYBYL, Pentacle&#x2a;</td>
<td align="left">Fractional factorial design</td>
<td align="right">0.75</td>
<td align="left">Yes</td>
<td align="left">Yes</td>
<td align="left">&#x2a;&#x2a;</td>
</tr>
<tr>
<td align="left">Multiple linear regression (MLR) (<xref ref-type="bibr" rid="B42">Xu et al., 2015</xref>)</td>
<td align="right">0.943</td>
<td align="left">ISIS/Draw, CODESSA&#x2a;</td>
<td align="left">Forward selection</td>
<td align="right">0.848</td>
<td align="left"/>
<td align="left"/>
<td align="left">&#x2a;&#x2a;</td>
</tr>
<tr>
<td align="left">MLR (<xref ref-type="bibr" rid="B25">Mirrahimi et al., 2016</xref>)</td>
<td align="right">0.841</td>
<td align="left">ISIS/Draw, MOPAC, Web-DRAGON&#x2a;</td>
<td align="left">Genetic algorithm</td>
<td align="right">0.821</td>
<td align="left">Yes</td>
<td align="left">Yes</td>
<td align="left">&#x2a;&#x2a;</td>
</tr>
<tr>
<td align="left">MLR (<xref ref-type="bibr" rid="B28">Prakasvudhisarn et al., 2009</xref>)</td>
<td align="right">0.833</td>
<td align="left">HyperChem, DRAGON&#x2a;</td>
<td align="left">Forward selection</td>
<td align="right">0.826</td>
<td align="left">Yes</td>
<td align="left">Yes</td>
<td align="left">&#x2a;&#x2a;</td>
</tr>
<tr>
<td align="left">MLR (<xref ref-type="bibr" rid="B36">Tropsha, 2010</xref>)</td>
<td align="right">0.78</td>
<td align="left">SYBYL, MOE, AutoDock Tools, BINANA</td>
<td align="left">Genetic algorithm</td>
<td align="right">0.82</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">Yes</td>
<td align="left">&#x2a;&#x2a;</td>
</tr>
<tr>
<td align="left">Artificial neural network (<xref ref-type="bibr" rid="B28">Prakasvudhisarn et al., 2009</xref>)</td>
<td align="right">0.957</td>
<td align="left">HyperChem, DRAGON&#x2a;</td>
<td align="left">Forward selection</td>
<td align="right">0.955</td>
<td align="left">Yes</td>
<td align="left">Yes</td>
<td align="left">&#x2a;&#x2a;</td>
</tr>
<tr>
<td align="left">Support vector machine (<xref ref-type="bibr" rid="B37">Trott and Olson, 2010</xref>)</td>
<td align="right">0.971</td>
<td align="left">HyperChem, MOE&#x2a;</td>
<td align="left">Particle swarm</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
<td align="left">&#x2a;&#x2a;</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>&#x2a;Presence of a paywall, usually due to specialized software that requires a license.</p>
</fn>
<fn>
<p>&#x2a;&#x2a;Insufficient verification. None of the models investigated were both fully validated and openly accessible.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>In this study, the accuracy and usability of docking and QSARs were compared in order to establish an appropriate framework for the computational design of cyclodextrin based affinity delivery devices. Autodock VINA, an open-source docking program developed by Trott, was used to investigate docking methods (<xref ref-type="bibr" rid="B37">Trott and Olson, 2010</xref>). A variety of statistical methods presented in previous cyclodextrin QSARs were also investigated. The performance of QSARs was evaluated on both a standard test set as well as an external validation set to confirm accuracy. Properly evaluating the use of docking and QSARs should improve selection of possible guests for cyclodextrin, reducing the rejection of good candidates (Type II error) and limiting experimental investigation of bad candidates (Type I error).</p>
<p>Finally, though the coded models could be made freely available, understanding the raw script remained a significant obstacle for new users. Additionally, users would have to download multiple files and programs to their own computers, creating potential issues with device compatibility, storage restrictions, processor limitations, etc. To overcome these obstacles and improve accessibility, the models were then integrated into a web application built with the R library &#x201c;shiny&#x201d; and then uploaded online. To demonstrate the ease of extendability of the app and its value in planning drug delivery strategies the results from the QSAR studies were then integrated into a mechanistic model of drug release.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>2 Materials and methods</title>
<p>In order to be accessible, the models use only open-source software. Importing experimental data, cleaning data, and creating QSAR models were performed using R in RStudio. Both the coding language and the IDE are freely downloadable and easily accessible on Windows, Mac OS, and Linux. Descriptors were generated with PaDEL, also freely downloadable and open-source. Only the original observations of cyclodextrin complexation energies remain inaccessible to the public, but this does not have any effect on using the models for new predictions.</p>
<sec id="s2-1">
<title>2.1 Dataset</title>
<p>Many of the models in <xref ref-type="table" rid="T1">Table 1</xref> work from the same data source, a compilation of &#x3b1;- and &#x3b2;-CD affinities published by Suzuki in 2001 (<xref ref-type="bibr" rid="B35">Suzuki, 2001</xref>) [additionally, the sources that cite a different paper by Katritzky ultimately use this same data, as the Katritzky paper cites Suzuki for observations (<xref ref-type="bibr" rid="B17">Katritzky et al., 2004</xref>)]. In addition to Suzuki, we also compiled complexes of &#x3b1;- and &#x3b2;-CD Rekharsky and Inoue and Suzuki (<xref ref-type="bibr" rid="B30">Rekharsky and Inoue, 1998</xref>). Complexes of &#x3b3;-CD, missing from the Suzuki dataset and sparse in the Rekharsky and Inoue data, were collected from Connors (<xref ref-type="bibr" rid="B4">Connors, 1995</xref>). Once compiled, the data were cleaned for reliable information, one-to-one cyclodextrin complexes, a temperature of 298 &#xb1; 2&#xa0;K, and a solvent of water with pH 7. To obtain structure-data files (SDFs) of the ligands, the names of the guest molecules were passed through the Chemical Identifier Resolver, a web interface provided by the National Cancer Institute&#x2019;s Computer-Aided Drug Design Group (NCI/CADD). To handle the data, the R packages tidyverse, data. table, XML, RCurl, and Matrix were used (<xref ref-type="bibr" rid="B2">Bates et al., 2017</xref>; <xref ref-type="bibr" rid="B8">Dowle et al., 2017</xref>; <xref ref-type="bibr" rid="B41">Wickham, 2017</xref>; <xref ref-type="bibr" rid="B9">Duncan Temple Lang and the CRAN T and eam, 2018a</xref>; <xref ref-type="bibr" rid="B10">Duncan Temple Lang and the CRAN Team, 2018b</xref>).</p>
<p>Dataset splitting was performed using the R package caret (<xref ref-type="bibr" rid="B18">Kuhn and Quinlan, 2018</xref>). First, the cleaned data was split between &#x3b1;-, &#x3b2;-, and &#x3b3;-CD. Structural and activity outliers in each category were removed. Structural outliers were detected using a statistical method relying on standard deviations of molecular descriptors (<xref ref-type="bibr" rid="B32">Roy et al., 2015</xref>). For activity outliers, molecules with reported &#x394;G values greater than 2.5 standard deviations from the mean were removed. Though traditional practice advises classifies outliers as values more than only two standard deviations away, in this case, retaining data points remained a priority and a larger margin was allowed. There were 9, 21, and 11 &#x3b1;-, &#x3b2;-, and &#x3b3;-CD outliers, respectively. After removal, around 200, 250, and 100 &#x3b1;-, &#x3b2;-, and &#x3b3;-CD observations remained.</p>
<p>The data was then split into training, testing data and external validation. For each separate cyclodextrin, an external validation set was created from a random 15% subset of the data. To create multiple training and test sets, the remaining modeling data was split with representative resampling of &#x394;G values into ten different 75:25 train to test data partitions. Though not as advanced as maximum dissimilarity algorithms, this method proved more practical due to the large number of descriptors (over 1,000) generated for each guest molecule. Furthermore, maximum dissimilarity algorithms, when implemented in this instance, had the unfortunate tendency to select highly similar training and test sets, defeating the purpose of creating multiple sets in the first place.</p>
</sec>
<sec id="s2-2">
<title>2.2 Docking calculations</title>
<p>The process of docking is based on two processes: sampling and scoring (<xref ref-type="bibr" rid="B16">Jacob et al., 2012</xref>). Sampling refers to the capacity to search an active site on a protein, macromolecule or, in this case, affinity host. This can be performed with distance matrices, matching algorithms or incremental construction, multiple copy simultaneous searching, stochastic methods, or any combination of the aforementioned strategies. Scoring calculates the final binding affinity between the guest and host and can be dependent on force-field, empirical, or knowledge-based calculations. Docking generally involves the use of a host and a guest molecule which can be either rigid or flexible. Three types of conformation exist: rigid-rigid, rigid-flexible and flexible-flexible. In this paper we use AutoDock Vina, a version of AutoDock that uses Monte Carlo stochastic sampling coupled with a force field based scoring function from a resample of a drug like database to derive its weighted parameters. Vina in particular uses a flexible drug guest and a rigid cyclodextrin host, although it allows side chain mobility when docking ligands onto proteins.</p>
<p>The PyRx Virtual Screening Tool provides a variety of services, including molecular energy minimization, docking calculation, and visualization of molecules. PyRx version 0.8 was used here, as further editions require purchase (<xref ref-type="bibr" rid="B6">Dallakyan and J, 2015</xref>). Ostensibly, the source code of newer versions of PyRx is freely available, but actually implementing the code requires fairly advanced knowledge of Python, making public usage difficult. To begin, all guest molecules went through energy minimization to determine the most likely atomic configurations. AutoDock Vina, integrated within PyRx, calculated the change in Gibbs free energy (kcal/mol). We tested the effect on the docking process of changes in the search space, search exhaustiveness, and scoring force field type.</p>
</sec>
<sec id="s2-3">
<title>2.3 Descriptor generation</title>
<p>The open source software PaDEL-Descriptor calculated over 1,000 descriptors for the remaining molecules, including fingerprints, structural details, and physical properties (<xref ref-type="bibr" rid="B43">Yap, 2011</xref>). Additionally, PaDEL-Descriptor removed salts and minimized the energy of inputted files using an MM2 force field. To improve model interpretability, more abstract predictors, such as those related to eigenvalues for molecular matrices or autocorrelation, were excluded from calculation. The elimination of these descriptors did not produce any noticeable effect on final model accuracy and made feature selection less resource intensive.</p>
</sec>
<sec id="s2-4">
<title>2.4 Feature selection</title>
<p>Recursive feature elimination (RFE), implemented with caret, was used to subset the predictors used for model-building (Kuhn 2018). Using this method, a random forest model is created using all available descriptors. Once trained, the relative importances of the predictors are calculated and differently sized subsets (defined by the user) of variables are selected to create and evaluate new models. The best combination of predictors is then returned by the model. RFE was performed on each of the ten train-test splits. The predictors determined to be useful for all folds were saved and used for tuning and training the models. This resulted in 13 variables for &#x237a;-CD, 16 variables for &#x3b2;-CD, and 39 variables for &#x3b3;-CD.</p>
</sec>
<sec id="s2-5">
<title>2.5 QSAR development</title>
<p>We investigated the accuracy of several models that appeared in previous attempts at cyclodextrin QSARS (<xref ref-type="table" rid="T1">Table 1</xref>), including Cubist models, generalized linear models (GLM or GLMNet), random forests, partial least squares models, and support vector machines. Additionally, two QSAR methods not previously published for cyclodextrin&#x2014;multivariate adaptive regression splines (MARS) and gradient-boosted models&#x2014;were created and evaluated. Model building was accomplished with R-packages Cubist, glmnet, randomForest, pls, e1071, earth, and gbm, respectively (<xref ref-type="bibr" rid="B5">Cutler and Wiener, 2015</xref>; <xref ref-type="bibr" rid="B23">Mevik and Liland, 2016</xref>; <xref ref-type="bibr" rid="B12">Friedman et al., 2017</xref>; <xref ref-type="bibr" rid="B19">Kuhn et al., 2017</xref>; <xref ref-type="bibr" rid="B24">Meyer et al., 2017</xref>).</p>
<p>Cross-validation was used to determine ideal tuning parameters for each QSAR. For faster QSARs&#x2014;such as generalized linear models (GLM) and partial least squares (PLS)&#x2014;tuning was performed using 10-fold cross validation. For more resource-intensive models or models with large parameter spaces&#x2014;such as random forests, Cubist and support vector machines (SVM)&#x2013; only five folds were used. Optimized models, QSARs built with the tuned parameters and trained on the entire training set, were used to predict the test for each combination of test and training set. Further fine tuning was also performed at this step. The model that produced the lowest root-mean square error (RMSE) and highest <italic>R</italic>
<sup>2</sup> (or an otherwise most ideal combination) on all the test sets became the final model, i.e., the model saved for future use. Furthermore, the models were evaluated according to Tropsha and Golbraikh standards for QSARs (<xref ref-type="bibr" rid="B15">Golbraikh and Tropsha, 2002</xref>). Although <italic>R</italic>
<sup>2</sup> and RMSE can be useful for generalizing predictive capacity, they may be misleading in certain cases, necessitating stricter additional standards of evaluation. As an additional test of reproducibility, the final models were used in ensemble to predict the values of the external validation set. Because this dataset was withheld from the entire model training process, the external validation set served to simulate model performance on new data.</p>
</sec>
<sec id="s2-6">
<title>2.6 Applicability domain</title>
<p>Applicability domain describes the range of molecules where the model can be expected to generate reliable predictions. A new molecule outside of the applicability domain is structurally quite different from the set of data the model was trained on, and thus a prediction will rely on extrapolation and may not be accurate. The applicability domain of the models was determined with the same method used to detect outliers when cleaning the dataset (<xref ref-type="bibr" rid="B32">Roy et al., 2015</xref>).</p>
</sec>
<sec id="s2-7">
<title>2.7 Y-randomization</title>
<p>Y-randomization was used to further verify the significance of the results. Many advanced QSAR methods are powerful enough to model data off of noise, so y-randomization ensures that the modelling process produces results significantly more accurate than what could be obtained by chance. Randomization can be achieved by permutation (randomly changing the positions of observed values) or random number generation (replacing observed values with completely new data). Different combinations of permutation and/or random generation yields five different modes of y-randomization to investigate: 1) original &#x394;G values vs randomly generated descriptors, 2) permuted &#x394;G vs original descriptors, 3) random &#x394;G vs original descriptors, 4) random &#x394;G vs random descriptors, and 5) permuted &#x394;G vs random descriptors. (Combinations including permutation of descriptors are not included because the large number of predictors in QSARs renders the effects of such a process virtually indistinguishable from random number generation.) However, because the y-randomization process is extremely resource-intensive (as each mode requires that several randomized iterations undergo the modeling process), only mode 1, the most common interpretation of y-randomization, was investigated (<xref ref-type="bibr" rid="B33">R&#xfc;cker et al., 2007</xref>). The observed &#x394;G values were randomly assigned to guest molecules, and the entire model refitting process was re-done, from feature selection to external validation.</p>
</sec>
<sec id="s2-8">
<title>2.8 Creating an app</title>
<p>Using R&#x2019;s &#x201c;shiny&#x201d; package, most of the process of running the QSAR could be implemented in a web app. The app was split into three main pages: Download, Upload, and Explore. &#x201c;Download&#x201d; accesses Chemical Identifier Resolver and obtain SDFs. The page also draws the obtained molecule using the package &#x201c;ChemmineR,&#x201d; allowing the user to check that the SDF is accurate. &#x201c;Upload&#x201d; implements the QSARs after the user provides the app with a CSV of the descriptors from PaDEL-descriptor. After calculating the affinity and analyzing the applicability domain of the molecules, the user is provided with both a graph and a table of the results. The third page &#x201c;Explore,&#x201d; stores the results of using the ensemble on FDA-approved drugs, as obtained from the annual publication &#x201c;Orange Book: Approved Drug Products with Therapeutic Equivalence Evaluations.&#x201d; (<xref ref-type="bibr" rid="B11">Food and Drug Administration, 2019</xref>)</p>
</sec>
<sec id="s2-9">
<title>2.9 Modeling release curves</title>
<p>Partial differential equations that model drug diffusion were solved using the R package deSolve (<xref ref-type="bibr" rid="B34">Soetaert et al., 2018</xref>) using the method of lines as previously done by <xref ref-type="bibr" rid="B13">Fu et al. (2011)</xref>. In the model, the release media was assumed to be water and the delivery system was assumed to be flat, thin circular cyclodextrin disc. The boundary condition between the polymer and the media was approached as described by <xref ref-type="bibr" rid="B40">Wang and von Recum (2011)</xref>. Diffusivities of drug molecules were calculated from molecular weight and viscosity using a modified Stokes-Einstein-Sutherland equation, as done by <xref ref-type="bibr" rid="B39">Vulic et al. (2015)</xref>.</p>
</sec>
</sec>
<sec sec-type="results|discussion" id="s3">
<title>3 Results and discussion</title>
<sec id="s3-1">
<title>3.1 Performance of docking</title>
<p>When predicting on the entire cleaned dataset (all modeling data, which includes the training, testing, and external validation set), AutoDock Vina yielded an <italic>R</italic>
<sup>2</sup> of 0.18 and a RMSE of 5.00&#xa0;kJ/mol (<xref ref-type="fig" rid="F2">Figure 2</xref>). Of the 547 cleaned complexes, docking provided calculations for 458, failing to provide data on 89 complexes. Adjusting settings in Vina, such as the minimization algorithm size of the steps in the calculation, did not yield significant differences in accuracy. In comparison, the affinity of only around 40 cleaned molecules could not be obtained by the ensemble QSAR. In these cases, the withheld molecules were determined to be outliers, and the actual QSAR model could still be used to predict a value.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Results of PyRx docking.</p>
</caption>
<graphic xlink:href="frsfm-04-1402702-g002.tif"/>
</fig>
</sec>
<sec id="s3-2">
<title>3.2 Performance of QSARs</title>
<p>The results of predicting on the test data for each QSAR and cyclodextrin type are markedly higher than docking (with the exception of &#x3b3;-CD), reaching an <italic>R</italic>
<sup>2</sup> of around 0.5 to 0.7, as seen in <xref ref-type="table" rid="T2">Table 2</xref> and <xref ref-type="fig" rid="F3">Figure 3</xref>. The reported <italic>R</italic>
<sup>2</sup> for each QSAR type is calculated from an average of the performance of the model on all test splits. Additionally, <xref ref-type="table" rid="T2">Table 2</xref> contains information on verification of all the QSAR types (3-VII Validation methods). Of the types investigated, only PLS and GLMNet failed to pass the salvo of verification criteria, both falling short of attaining an <italic>R</italic>
<sup>2</sup> of 0.6.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Evaluation of QSARs on test sets.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="2" align="left">QSAR</th>
<th colspan="4" align="left">&#x3b1;-CD</th>
<th colspan="4" align="left">&#x3b2;-CD</th>
<th colspan="4" align="left">&#x3b3;-CD</th>
</tr>
<tr>
<th align="left">
<italic>A</italic>
</th>
<th align="left">
<italic>B</italic>
</th>
<th align="left">
<italic>C</italic>
</th>
<th align="left">
<italic>D</italic>
</th>
<th align="left">
<italic>A</italic>
</th>
<th align="left">
<italic>B</italic>
</th>
<th align="left">
<italic>C</italic>
</th>
<th align="left">
<italic>D</italic>
</th>
<th align="left">
<italic>A</italic>
</th>
<th align="left">
<italic>B</italic>
</th>
<th align="left">
<italic>C</italic>
</th>
<th align="left">
<italic>D</italic>
</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Cubist</td>
<td align="left">0.63</td>
<td align="left">0.56</td>
<td align="left">0.07</td>
<td align="left">0.93</td>
<td align="left">0.75</td>
<td align="left">0.59</td>
<td align="left">0.02</td>
<td align="left">0.97</td>
<td align="left">0.08&#x2a;</td>
<td align="left">&#x2212;0.19&#x2a;</td>
<td align="left">0</td>
<td align="left">0.98</td>
</tr>
<tr>
<td align="left">GBM</td>
<td align="left">0.78</td>
<td align="left">0.5</td>
<td align="left">0.05</td>
<td align="left">0.95</td>
<td align="left">0.83</td>
<td align="left">0.73</td>
<td align="left">0.02</td>
<td align="left">0.98</td>
<td align="left">0.35&#x2a;</td>
<td align="left">0.03&#x2a;</td>
<td align="left">0.68&#x2a;</td>
<td align="left">0.96</td>
</tr>
<tr>
<td align="left">GLMNet</td>
<td align="left">0.53&#x2a;</td>
<td align="left">0.54</td>
<td align="left">0</td>
<td align="left">0.94</td>
<td align="left">0.52&#x2a;</td>
<td align="left">0.45</td>
<td align="left">0.03</td>
<td align="left">0.95</td>
<td align="left">0.36&#x2a;</td>
<td align="left">&#x2212;0.26&#x2a;</td>
<td align="left">0.17&#x2a;</td>
<td align="left">0.97</td>
</tr>
<tr>
<td align="left">MARS</td>
<td align="left">0.65</td>
<td align="left">0.58</td>
<td align="left">0.03</td>
<td align="left">0.99</td>
<td align="left">0.73</td>
<td align="left">0.58</td>
<td align="left">0</td>
<td align="left">0.98</td>
<td align="left">0.40&#x2a;</td>
<td align="left">&#x2212;0.33&#x2a;</td>
<td align="left">0.33&#x2a;</td>
<td align="left">0.97</td>
</tr>
<tr>
<td align="left">PLS</td>
<td align="left">0.55&#x2a;</td>
<td align="left">0.47</td>
<td align="left">0.05</td>
<td align="left">0.93</td>
<td align="left">0.55&#x2a;</td>
<td align="left">0.47</td>
<td align="left">0.02</td>
<td align="left">0.95</td>
<td align="left">0.16&#x2a;</td>
<td align="left">&#x2212;0.1&#x2a;</td>
<td align="left">0.28&#x2a;</td>
<td align="left">0.97</td>
</tr>
<tr>
<td align="left">Polynomial SVM</td>
<td align="left">0.65</td>
<td align="left">0.55</td>
<td align="left">0</td>
<td align="left">0.98</td>
<td align="left">0.74</td>
<td align="left">0.56</td>
<td align="left">0</td>
<td align="left">0.98</td>
<td align="left">0.45&#x2a;</td>
<td align="left">&#x2212;0.28&#x2a;</td>
<td align="left">0.05</td>
<td align="left">1</td>
</tr>
<tr>
<td align="left">Random Forest</td>
<td align="left">0.76</td>
<td align="left">0.63</td>
<td align="left">0.02</td>
<td align="left">0.96</td>
<td align="left">0.84</td>
<td align="left">0.67</td>
<td align="left">0.03</td>
<td align="left">0.98</td>
<td align="left">0.69</td>
<td align="left">0.28&#x2a;</td>
<td align="left">0.24&#x2a;</td>
<td align="left">0.98</td>
</tr>
<tr>
<td align="left">RBF SVM</td>
<td align="left">0.74</td>
<td align="left">0.64</td>
<td align="left">0.02</td>
<td align="left">0.96</td>
<td align="left">0.85</td>
<td align="left">0.61</td>
<td align="left">0</td>
<td align="left">0.98</td>
<td align="left">0.35&#x2a;</td>
<td align="left">&#x2212;0.19&#x2a;</td>
<td align="left">0.05</td>
<td align="left">0.98</td>
</tr>
<tr>
<td align="left">Sigmoid SVM</td>
<td align="left">0.51&#x2a;</td>
<td align="left">0.52</td>
<td align="left">0.05</td>
<td align="left">0.93</td>
<td align="left">0.50&#x2a;</td>
<td align="left">0.56</td>
<td align="left">0.27</td>
<td align="left">0.92</td>
<td align="left">0.32&#x2a;</td>
<td align="left">&#x2212;0.11&#x2a;</td>
<td align="left">0</td>
<td align="left">0.98</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>The columns labeled A-D indicate the four conditions outlined by Golbraikh and Tropsha. A: <italic>R</italic>
<sup>2</sup> &#x3e; 0.6; B: q<sup>2</sup> &#x3e; 0.5, where q<sup>2</sup> is the result from leave one out cross-validation on the training set; C: &#x7c;<italic>R</italic>
<sup>2</sup>&#x2014;R&#x2032;<sup>2</sup>
<sub>0</sub>&#x7c;/<italic>R</italic>
<sup>2</sup> &#x3c; 0.1, indicating that the <italic>R</italic>
<sup>2</sup> when the axes are flipped (R<sup>&#x2019;2</sup>
<sub>0</sub>) is close to the original <italic>R</italic>
<sup>2</sup>; D: 0.85 &#x3c; k &#x3c; 1.15, where k, the slope of the regression line through the points is close to 1.</p>
</fn>
<fn>
<p>&#x2a;Model failed condition.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Results of QSARs on test sets.</p>
</caption>
<graphic xlink:href="frsfm-04-1402702-g003.tif"/>
</fig>
<p>In terms of reliability, most models were able to handle the available data well, providing calculations for all provided molecules. Only the Random Forest and Cubist models failed to calculate the affinity of some molecules, possibly due to being based around decision-trees. The algorithm underlying both models attempts to draw predictions by categorizing entries based on their features. If they encounter a molecules entirely different from the data they trained on, the models may fail to create a prediction. Advantageously for our approach, the failure to calculate some values becomes less important where models are combined in an ensemble where the final prediction is averaged over many models.</p>
<p>The results of ensemble prediction (averaging the results of many different QSARs) can be seen in <xref ref-type="fig" rid="F4">Figure 4</xref>. While &#x237a;- and &#x3b2;-CD models managed to reach moderately high predictive performance metrics, unfortunately, all &#x3b3;-CD models lacked useable predictive power.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>QSAR ensemble prediction.</p>
</caption>
<graphic xlink:href="frsfm-04-1402702-g004.tif"/>
</fig>
<p>The models passing the verification in <xref ref-type="table" rid="T2">Table 2</xref> were further verified using y-randomization. To ensure accuracy was not the result of the models building off of noise, 25 different permutations of &#x394;G values were created. All Q<sup>2</sup> values of the models created from the original data were calculated to lie well outside 3 standard deviations of the mean Q<sup>2</sup> of the randomized data. Additionally, the <italic>R</italic>
<sup>2</sup> values of the ensemble QSARs were significantly greater than the <italic>R</italic>
<sup>2</sup> values obtained from the ensemble models created from permuted data (means of 0.021 and 0.027 and standard deviations of 0.011 and 0.016 for &#x237a;- and &#x3b2;-CD, respectively).</p>
</sec>
<sec id="s3-3">
<title>3.3 Variable importance</title>
<p>Interpretability of a model provides a rough check if a model is calculating off of random noise or if the model is drawing logical calculations from physical properties to molecular behavior. Each model, due to differences in statistical algorithms and approaches, has differing levels of interpretability. GLM, being similar to linear models, have easily accessible coefficients associated with each predictor, so the relative impact of each factor can be compared with reasonable confidence. Cubist models, on the other hand, tend to be difficult to interpret as variables are processed through multiple levels of decision trees.</p>
<p>Evaluation for the relative importance of variables are shown in <xref ref-type="fig" rid="F5">Figure 5</xref>. Random forest was the only QSAR type with a pre-packaged importance function for variable analysis. PLS variables were analyzed using a function obtainable from <xref ref-type="bibr" rid="B22">Mevik et al. (2007)</xref>. Max Kuhn&#x2019;s caret package was used to evaluate GLMNet, the two SVM kernels, and Cubist. Unfortunately, caret was unable to process the final models for GLMNet and SVM, and the reported variable importance values were actually derived from models created within caret&#x2019;s &#x201c;train&#x201d; function, and are thus slightly different from the models saved in the ensemble. To determine importance, the &#x201c;train&#x201d; function removes a variable, rebuilds the model, and analyzes the effect on accuracy. The more important a variable, the larger the drop in accuracy. After each variable has been tested, the function can then rank the importance of the descriptors.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Variable importance.</p>
</caption>
<graphic xlink:href="frsfm-04-1402702-g005.tif"/>
</fig>
<p>For &#x3b2;-CD, XLogP, a measure of lipophilicity, appears to be important for all models, consistent with how the structure of cyclodextrin allows for easier complexation with small hydrophobic drugs. The same reasoning can be extended to LipoAffinityIndex and MLogP, additional approaches to quantifying lipophilicity. The number of carbons, nC, is also consistently important, possibly due to a relationship with molecule size. WTPT-2 is the PaDEL weighted path descriptor divided by the number of atoms, and also may be important due to encoding information on molecular size.</p>
<p>However, not all the variables can be linked to set chemical properties. SpMax and SpMin relates to eigenvalues of a modified connectivity matrix, a numerical representation of atomic and molecular bonds, and may not be associated with any interpretable physical property (the same analysis can also be used for GATS predictors). To aid interpretability, building a model with predictors easily attributed to physical or chemical properties may be advised. The extent to which interpretability should trade off with accuracy remains in question. Our findings go in accordance to those in the general literature were lipophilicity tends to highly influence model output.</p>
</sec>
<sec id="s3-4">
<title>3.4 Web application and FDA database</title>
<p>After collecting a list of FDA-approved drugs and drug combinations from the Orange Book, an annual publication listing all approved pharmaceuticals, the names were cleaned for individual active compounds. In total, 1,401 unique molecules could be extracted. Of these, 1,116 could be downloaded from Cactus and 1,031 could be processed by PaDEL. Many of the molecules that could not be analyzed by PaDEL would have proven impractical for cyclodextrin delivery, such as simple ionic salts (e.g., potassium chloride), or large molecules made of more than 100 atoms. Running the remaining guests through applicability domain analysis yielded 638 molecules, 45.5% of the original set. While less than half of FDA-approved drugs could pass through the model, the 600 available guests spans a wide range of properties and uses, allowing the page to be useful for candidate selection (<xref ref-type="fig" rid="F6">Figure 6</xref>).</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>User interface of the Shiny App.</p>
</caption>
<graphic xlink:href="frsfm-04-1402702-g006.tif"/>
</fig>
<p>Though the app could be uploaded online through shinyapps. io, server time limitations on the account hosting the app make it impractical for usage by a large number of individuals simultaneously. In order to run the app for more than a few hours, such as with screening a large dataset of molecules, the code would have to be downloaded through GitHub. In addition, the user would need to download the R libraries and the IDE RStudio, potentially negating the goal of creating an accessible, intuitive interface. The &#x201c;Explore&#x201d; page partially alleviates this obstacle, as it allows the user to perform a quick search of a pre-predicted affinity rather than spend time downloading the structure file, launching PaDEL, and running the QSAR.</p>
</sec>
<sec id="s3-5">
<title>3.5 Drug release module</title>
<p>To demonstrate the extensibility of the shiny app and its value in the design of drug delivery strategies the results of the QSAR predictions can be fed into a mechanistic model of drug delivery (<xref ref-type="fig" rid="F7">Figure 7</xref>). The results demonstrate the ranges of values expected from the strongest affinity binding predictions and from the weakest. As expected, strong predictions produce much slower release profiles and weak predictions produce faster release profiles. Conservation of mass was verified as the sum of all mass within the system from the polymer and media compartment across all times as a test of the implementation. Notably, the implementation in R required a modification of the method of lines for appropriate modeling of the polymer to liquid media interface. (<xref ref-type="bibr" rid="B20">Linge and Langtangen, 2016</xref>). Future efforts in creating new modules could explore substructure searching to identify alternative strategies for weak binders or drugs that demonstrate unsuitable release profiles. It is expected that not all drugs will follow this simplistic model of drug release. For those cases our lab has built a whole suite of approaches to alter elution rates such as a wide range of formulations, supramolecular interactions, Schiff-base formation and multi-arm PEG substitutions.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Drug release curves.</p>
</caption>
<graphic xlink:href="frsfm-04-1402702-g007.tif"/>
</fig>
</sec>
</sec>
<sec sec-type="conclusion" id="s4">
<title>4 Conclusion</title>
<p>In predicting the binding affinity of cyclodextrin with small drug molecules, QSARS such as Cubist, GBM, MARS, random forest, and SVM models can be created using accessible open-source software. These models outperform available docking software in both accuracy and time consumption and pass statistical verification of reliability. The additional accuracy afforded by QSARs can be integrated into the previously published mechanistic model for predicting drug release curves for candidate molecules. This would both help narrow down candidates for cyclodextrin affinity-based drug delivery as well as help advise which molecules are most appropriate to tailor the release rate from a delivery system for a given biomedical application. Furthermore, the QSAR models can be used to evaluate existing marketed pharmaceutical formulations for their small molecule interaction with cyclodextrin to better understand the extent that the strength of binding between the cyclodextrin and the drug is of importance for the marketed product formulation (<xref ref-type="bibr" rid="B3">Braga, 2023</xref>; <xref ref-type="bibr" rid="B29">Pusk&#xe1;s et al., 2023</xref>). Both of these goals can be achieved by any reader interested in the current work by accessing the github repository for this manuscript (<ext-link ext-link-type="uri" xlink:href="https://github.com/awqx/qsar-app">https://github.com/awqx/qsar-app</ext-link>). The current model is limited to predictions in the experimental space of the training data and applications outside its applicability, for example, at low or very high pH, should be employed with caution and tested experimentally.</p>
<p>The integration of these machine learning models in combination with the mechanistic models of drug delivery all within a web application allows for a novel framework to plan drug delivery strategies. The application allows for a &#x201c;design before you build&#x201d; approach where others can bring their library of small molecules and determine which ones make the best candidates for an affinity release strategy. The mechanistic models of other geometries or drug delivery forms such as microparticles and injectable polymers can be readily included in the application to further extend the capabilities for other biomedical applications.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The datasets presented in this study can be found in online repositories. The names of the repository/repositories and accession number(s) can be found below: QSAR model building package repository <ext-link ext-link-type="uri" xlink:href="https://github.com/awqx/qsarr">https://github.com/awqx/qsarr</ext-link> Drug Release Module Code <ext-link ext-link-type="uri" xlink:href="https://github.com/eriveradelgado/ODE_Practice/blob/master/09_ODE-drug-release.Rmd">https://github.com/eriveradelgado/ODE_Practice/blob/master/09_ODE-drug-release.Rmd</ext-link> QSAR Application <ext-link ext-link-type="uri" xlink:href="https://github.com/awqx/qsar-app">https://github.com/awqx/qsar-app</ext-link> Walkthrough on how to use the models <ext-link ext-link-type="uri" xlink:href="https://github.com/awqx/qsar-app">https://github.com/awqx/qsar-app</ext-link> Enter subfile process. Rmd.</p>
</sec>
<sec id="s6">
<title>Author contributions</title>
<p>AX: Investigation, Software, Writing&#x2014;original draft, Data curation. ER-D: Conceptualization, Supervision, Investigation, Software, Writing&#x2014;original draft. HR: Conceptualization, Funding acquisition, Supervision, Writing&#x2014;review and editing.</p>
</sec>
<sec sec-type="funding-information" id="s7">
<title>Funding</title>
<p>The author(s) declare that no financial support was received for the research, authorship, and/or publication of this article.</p>
</sec>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s9">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ahmadi</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Ghasemi</surname>
<given-names>J. B.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>3D-QSAR and docking studies of the stability constants of different guest molecules with beta-cyclodextrin</article-title>. <source>J. Incl. Phenom. Macrocycl. Chem.</source> <volume>79</volume>, <fpage>401</fpage>&#x2013;<lpage>413</lpage>. <pub-id pub-id-type="doi">10.1007/s10847-013-0363-5</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Bates</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Maechler</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Davis</surname>
<given-names>T. A.</given-names>
</name>
<name>
<surname>Amd</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Oehlschl&#xe4;gel</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Riedy</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2017</year>). <article-title>Matrix: sparse and dense matrix classes and methods</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/Matrix/index.html">https://cran.r-project.org/web/packages/Matrix/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Braga</surname>
<given-names>S. S.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Molecular mind games: the medicinal action of cyclodextrins in neurodegenerative diseases</article-title>. <source>Biomolecules</source> <volume>13</volume> (<issue>4</issue>), <fpage>666</fpage>. <pub-id pub-id-type="doi">10.3390/biom13040666</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Connors</surname>
<given-names>K. A.</given-names>
</name>
</person-group> (<year>1995</year>). <article-title>Population characteristics of cyclodextrin complex stabilities in aqueous solution</article-title>. <source>J. Pharm. Sci.</source> <volume>84</volume>, <fpage>843</fpage>&#x2013;<lpage>848</lpage>. <pub-id pub-id-type="doi">10.1002/jps.2600840712</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Cutler</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Wiener</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>randomForest: Breiman and Cutler&#x2019;s random forests for classification and regression</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/randomForest/index.html">https://cran.r-project.org/web/packages/randomForest/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dallakyan</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>J</surname>
<given-names>Olson A.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Small-molecule library screening by docking with PyRx</article-title>. <source>Methods Mol. Biol. (Clifton, NJ)</source> <volume>1263</volume>, <fpage>243</fpage>&#x2013;<lpage>250</lpage>. <pub-id pub-id-type="doi">10.1007/978-1-4939-2269-7_19</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Dehmer</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Varmuza</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Bonchev</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2012</year>). <source>Statistical modelling of molecular descriptors in QSAR/QSPR</source>. <publisher-name>John Wiley and Sons</publisher-name>.</citation>
</ref>
<ref id="B8">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Dowle</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Srinivasan</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Gorecki</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Short</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Lianoglou</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Antonyan</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>data.table: extension of &#x201c;data.frame&#x201d;</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/data.table/index.html">https://cran.r-project.org/web/packages/data.table/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B9">
<citation citation-type="book">
<collab>Duncan Temple Lang and the CRAN Team</collab> (<year>2018a</year>). <source>RCurl: general network (HTTP/FTP/&#x2026;) client interface for R</source>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/RCurl/index.html">https://cran.r-project.org/web/packages/RCurl/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B10">
<citation citation-type="book">
<collab>Duncan Temple Lang and the CRAN Team</collab> (<year>2018b</year>). <source>XML: tools for parsing and generating XML within R and S-plus</source>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/XML/index.html">https://cran.r-project.org/web/packages/XML/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B11">
<citation citation-type="book">
<collab>Food and Drug Administration</collab> (<year>2019</year>). <source>Approved drug products with therapeutic equivalence evaluations</source>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://www.fda.gov/media/71474/download">https://www.fda.gov/media/71474/download</ext-link> (Accessed May 31, 2019)</comment>.</citation>
</ref>
<ref id="B12">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Friedman</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Hastie</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Simon</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Qian</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Tibshirani</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Glmnet: lasso and elastic-net regularized generalized linear models</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/glmnet/index.html">https://cran.r-project.org/web/packages/glmnet/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fu</surname>
<given-names>A. S.</given-names>
</name>
<name>
<surname>Thatiparti</surname>
<given-names>T. R.</given-names>
</name>
<name>
<surname>Saidel</surname>
<given-names>G. M.</given-names>
</name>
<name>
<surname>von Recum</surname>
<given-names>H. A.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>Experimental studies and modeling of drug release from a tunable affinity-based drug delivery platform</article-title>. <source>Ann. Biomed. Eng.</source> <volume>39</volume>, <fpage>2466</fpage>&#x2013;<lpage>2475</lpage>. <pub-id pub-id-type="doi">10.1007/s10439-011-0336-z</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ghasemi</surname>
<given-names>J. B.</given-names>
</name>
<name>
<surname>Salahinejad</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Rofouei</surname>
<given-names>M. K.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>An alignment independent 3D-QSAR study for predicting the stability constants of structurally diverse compounds with &#x3b2;-cyclodextrin</article-title>. <source>J. Incl. Phenom. Macrocycl. Chem.</source> <volume>71</volume>, <fpage>195</fpage>&#x2013;<lpage>206</lpage>. <pub-id pub-id-type="doi">10.1007/s10847-011-9927-4</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Golbraikh</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Tropsha</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2002</year>). <article-title>Beware of q2!</article-title>. <source>J. Mol. Graph. Model.</source> <volume>20</volume>, <fpage>269</fpage>&#x2013;<lpage>276</lpage>. <pub-id pub-id-type="doi">10.1016/S1093-3263(01)00123-1</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jacob</surname>
<given-names>R. B.</given-names>
</name>
<name>
<surname>Andersen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>McDougal</surname>
<given-names>O. M.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Accessible high-throughput virtual screening molecular docking software for students and educators</article-title>. <source>PLOS Comput. Biol.</source> <volume>8</volume>, <fpage>e1002499</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pcbi.1002499</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Katritzky</surname>
<given-names>A. R.</given-names>
</name>
<name>
<surname>Fara</surname>
<given-names>D. C.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Karelson</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Suzuki</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Solov&#x2019;ev</surname>
<given-names>V. P.</given-names>
</name>
<etal/>
</person-group> (<year>2004</year>). <article-title>Quantitative Structure&#x2212;Property relationship modeling of <italic>&#x3b2;</italic>-cyclodextrin complexation free energies</article-title>. <source>J. Chem. Inf. Comput. Sci.</source> <volume>44</volume>, <fpage>529</fpage>&#x2013;<lpage>541</lpage>. <pub-id pub-id-type="doi">10.1021/ci034190j</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Kuhn</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Quinlan</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2018</year>). <source>Caret: classification and regression training</source>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://CRAN.R-project.org/package=caret">https://CRAN.R-project.org/package&#x3d;caret</ext-link>.</comment>
</citation>
</ref>
<ref id="B19">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Kuhn</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Steve</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Chris</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Nathan</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Quinlan</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Cubist: rule- and instance-based regression modeling</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/Cubist/index.html">https://cran.r-project.org/web/packages/Cubist/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B20">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Linge</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Langtangen</surname>
<given-names>H. P.</given-names>
</name>
</person-group> (<year>2016</year>) &#x201c;<article-title>Texts in computational science and engineering</article-title>,&#x201d; in <source>Programming for computations - Python</source>. <pub-id pub-id-type="doi">10.1007/978-3-319-32428-9</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Merzlikine</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Abramov</surname>
<given-names>Y. A.</given-names>
</name>
<name>
<surname>Kowsz</surname>
<given-names>S. J.</given-names>
</name>
<name>
<surname>Thomas</surname>
<given-names>V. H.</given-names>
</name>
<name>
<surname>Mano</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>Development of machine learning models of &#x3b2;-cyclodextrin and sulfobutylether-&#x3b2;-cyclodextrin complexation free energies</article-title>. <source>Int. J. Pharm.</source> <volume>418</volume>, <fpage>207</fpage>&#x2013;<lpage>216</lpage>. <pub-id pub-id-type="doi">10.1016/j.ijpharm.2011.03.065</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Mevik</surname>
<given-names>B.-H.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>VIP.R: implementation of VIP (variable importance in projection) (&#x2a;) for the &#x201c;pls&#x201d; package</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="http://mevik.net/work/software/VIP.R">http://mevik.net/work/software/VIP.R</ext-link>.</comment>
</citation>
</ref>
<ref id="B23">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Mevik</surname>
<given-names>B.-H.</given-names>
</name>
<name>
<surname>Liland</surname>
<given-names>R. W.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Pls: partial least squares and principal component regression</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/pls/index.html">https://cran.r-project.org/web/packages/pls/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B24">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Meyer</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Dimitriadou</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Hornik</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Weingessel</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Leisch</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2017</year>) <article-title>e1071: Misc Functions of the Department of Statistics, Probability Theory Group (Formerly: E1071) TU Wien</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://cran.r-project.org/web/packages/e1071/index.html">https://cran.r-project.org/web/packages/e1071/index.html</ext-link>.</comment>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mirrahimi</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Salahinejad</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Ghasemi</surname>
<given-names>J. B.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>QSPR approaches to elucidate the stability constants between &#x3b2;-cyclodextrin and some organic compounds: docking based 3D conformer</article-title>. <source>J. Mol. Liq.</source> <volume>219</volume>, <fpage>1036</fpage>&#x2013;<lpage>1043</lpage>. <pub-id pub-id-type="doi">10.1016/j.molliq.2016.04.037</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ortiz</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Fragoso</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>O&#x2019;Sullivan</surname>
<given-names>C. K.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>Amperometric detection of antibodies in serum: performance of self-assembled cyclodextrin/cellulose polymer interfaces as antigen carriers</article-title>. <source>Org. Biomol. Chem.</source> <volume>9</volume>, <fpage>4770</fpage>&#x2013;<lpage>4773</lpage>. <pub-id pub-id-type="doi">10.1039/C1OB05473B</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>P&#xe9;rez-Garrido</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Helguera</surname>
<given-names>A. M.</given-names>
</name>
<name>
<surname>Guill&#xe9;n</surname>
<given-names>A. A.</given-names>
</name>
<name>
<surname>Cordeiro</surname>
<given-names>MNDS</given-names>
</name>
<name>
<surname>Escudero</surname>
<given-names>A. G.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Convenient QSAR model for predicting the complexation of structurally diverse compounds with &#x3b2;-cyclodextrins</article-title>. <source>Bioorg. Med. Chem.</source> <volume>17</volume>, <fpage>896</fpage>&#x2013;<lpage>904</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmc.2008.11.040</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Prakasvudhisarn</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wolschann</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Lawtrakul</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Predicting complexation thermodynamic parameters of &#x3b2;-cyclodextrin with chiral guests by using swarm intelligence and support vector machines</article-title>. <source>Int. J. Mol. Sci.</source> <volume>10</volume>, <fpage>2107</fpage>&#x2013;<lpage>2121</lpage>. <pub-id pub-id-type="doi">10.3390/ijms10052107</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pusk&#xe1;s</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Szente</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Sz&#x151;cs</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Fenyvesi</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Recent list of cyclodextrin-containing drug products</article-title>. <source>Period. Polytech. Chem. Eng.</source> <volume>67</volume> (<issue>1</issue>), <fpage>11</fpage>&#x2013;<lpage>17</lpage>. <pub-id pub-id-type="doi">10.3311/ppch.21222</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rekharsky</surname>
<given-names>M. V.</given-names>
</name>
<name>
<surname>Inoue</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>1998</year>). <article-title>Complexation thermodynamics of cyclodextrins</article-title>. <source>Chem. Rev.</source> <volume>98</volume>, <fpage>1875</fpage>&#x2013;<lpage>1918</lpage>. <pub-id pub-id-type="doi">10.1021/cr970015o</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rivera-Delgado</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Ward</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>von Recum</surname>
<given-names>H. A.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Providing sustained transgene induction through affinity-based drug delivery</article-title>. <source>J. Biomed. Mater Res.</source> <volume>104</volume>, <fpage>1135</fpage>&#x2013;<lpage>1142</lpage>. <pub-id pub-id-type="doi">10.1002/jbm.a.35643</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Roy</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Kar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ambure</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>On a simple approach for determining applicability domain of QSAR models</article-title>. <source>Chemom. Intelligent Laboratory Syst.</source> <volume>145</volume>, <fpage>22</fpage>&#x2013;<lpage>29</lpage>. <pub-id pub-id-type="doi">10.1016/j.chemolab.2015.04.013</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>R&#xfc;cker</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>R&#xfc;cker</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Meringer</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>y-Randomization and its Variants in QSPR/QSAR</article-title>. <source>J. Chem. Inf. Model</source> <volume>47</volume>, <fpage>2345</fpage>&#x2013;<lpage>2357</lpage>. <pub-id pub-id-type="doi">10.1021/ci700157b</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Soetaert</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Petzoldt</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Setzer</surname>
<given-names>R. W.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>deSolve: solvers for initial value problems of differential equations (&#x201c;ODE&#x201d;, &#x201c;DAE&#x201d;, &#x201c;DDE&#x201d;)</article-title>. <comment>Available at: <ext-link ext-link-type="uri" xlink:href="https://CRAN.R-project.org/package=deSolve">https://CRAN.R-project.org/package&#x3d;deSolve</ext-link>.</comment>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Suzuki</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2001</year>). <article-title>A nonlinear group contribution method for predicting the free energies of inclusion complexation of organic molecules with &#x3b1;- and &#x3b2;-cyclodextrins</article-title>. <source>J. Chem. Inf. Comput. Sci.</source> <volume>41</volume>, <fpage>1266</fpage>&#x2013;<lpage>1273</lpage>. <pub-id pub-id-type="doi">10.1021/ci010295f</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tropsha</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Best practices for QSAR model development, validation, and exploitation</article-title>. <source>Mol. Inf.</source> <volume>29</volume>, <fpage>476</fpage>&#x2013;<lpage>488</lpage>. <pub-id pub-id-type="doi">10.1002/minf.201000061</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Trott</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Olson</surname>
<given-names>A. J.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>AutoDock Vina: improving the speed and accuracy of docking with a new scoring function, efficient optimization, and multithreading</article-title>. <source>J. Comput. Chem.</source> <volume>31</volume>, <fpage>455</fpage>&#x2013;<lpage>461</lpage>. <pub-id pub-id-type="doi">10.1002/jcc.21334</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Veselinovi&#x107;</surname>
<given-names>A. M.</given-names>
</name>
<name>
<surname>Veselinovi&#x107;</surname>
<given-names>J. B.</given-names>
</name>
<name>
<surname>Toropov</surname>
<given-names>A. A.</given-names>
</name>
<name>
<surname>Toropova</surname>
<given-names>A. P.</given-names>
</name>
<name>
<surname>Nikoli&#x107;</surname>
<given-names>G. M.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>
<italic>In silico</italic> prediction of the &#x3b2;-cyclodextrin complexation based on Monte Carlo method</article-title>. <source>Int. J. Pharm.</source> <volume>495</volume>, <fpage>404</fpage>&#x2013;<lpage>409</lpage>. <pub-id pub-id-type="doi">10.1016/j.ijpharm.2015.08.078</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vulic</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Pakulska</surname>
<given-names>M. M.</given-names>
</name>
<name>
<surname>Sonthalia</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Ramachandran</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Shoichet</surname>
<given-names>M. S.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Mathematical model accurately predicts protein release from an affinity-based delivery system</article-title>. <source>J. Control. Release</source> <volume>197</volume>, <fpage>69</fpage>&#x2013;<lpage>77</lpage>. <pub-id pub-id-type="doi">10.1016/j.jconrel.2014.10.032</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>N. X.</given-names>
</name>
<name>
<surname>von Recum</surname>
<given-names>H. A.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>Affinity-based drug delivery</article-title>. <source>Macromol. Biosci.</source> <volume>11</volume>, <fpage>321</fpage>&#x2013;<lpage>332</lpage>. <pub-id pub-id-type="doi">10.1002/mabi.201000206</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Wickham</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2017</year>). <source>Tidyverse: easily install and load&#x2019;tidyverse&#x2019;packages. R. package version 1</source>.</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Wei</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Gu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Quantitative structure&#x2013;property relationship study of &#x3b2;-cyclodextrin complexation free energies of organic compounds</article-title>. <source>Chemom. Intelligent Laboratory Syst.</source> <volume>146</volume>, <fpage>313</fpage>&#x2013;<lpage>321</lpage>. <pub-id pub-id-type="doi">10.1016/j.chemolab.2015.06.001</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yap</surname>
<given-names>C. W.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>PaDEL-descriptor: an open source software to calculate molecular descriptors and fingerprints</article-title>. <source>J. Comput. Chem.</source> <volume>32</volume>, <fpage>1466</fpage>&#x2013;<lpage>1474</lpage>. <pub-id pub-id-type="doi">10.1002/jcc.21707</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>