<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Chem.</journal-id>
<journal-title>Frontiers in Chemistry</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Chem.</abbrev-journal-title>
<issn pub-type="epub">2296-2646</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1600945</article-id>
<article-id pub-id-type="doi">10.3389/fchem.2025.1600945</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Chemistry</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Two-dimensional QSAR-driven virtual screening for potential therapeutics against <italic>Trypanosoma cruzi</italic>
</article-title>
<alt-title alt-title-type="left-running-head">Maliyakkal et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fchem.2025.1600945">10.3389/fchem.2025.1600945</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Maliyakkal</surname>
<given-names>Naseer</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Kumar</surname>
<given-names>Sunil</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2813370/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name>
<surname>Bhowmik</surname>
<given-names>Ratul</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2387086/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Vishwakarma</surname>
<given-names>Harish Chandra</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Yadav</surname>
<given-names>Prabha</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Mathew</surname>
<given-names>Bijo</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2002543/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Department of Basic Medical Sciences</institution>, <institution>College of Applied Medical Sciences</institution>, <institution>King Khalid University</institution>, <addr-line>Khamis Mushait</addr-line>, <country>Saudi Arabia</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Department of Pharmaceutical Chemistry</institution>, <institution>Amrita School of Pharmacy</institution>, <institution>AIMS Health Sciences Campus</institution>, <institution>Amrita Vishwa Vidyapeetham</institution>, <addr-line>Kochi</addr-line>, <country>India</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Department of Pharmaceutical Chemistry</institution>, <institution>School of Pharmaceutical Education and Research</institution>, <institution>Jamia Hamdard</institution>, <addr-line>New Delhi</addr-line>, <country>India</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/832366/overview">Jorddy Neves Cruz</ext-link>, Federal University of Par&#xe1;, Brazil</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1466526/overview">Vijaya Bhaskar Baki</ext-link>, University of California, Riverside, United States</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2978700/overview">Li Ma</ext-link>, Xingimaging, United States</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Bijo Mathew, <email>bijomathew@aims.amrita.edu</email>, <email>bijovilaventgu@gmail.com</email>
</corresp>
<fn fn-type="equal" id="fn001">
<label>
<sup>&#x2020;</sup>
</label>
<p>These authors have contributed equally to this work</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>11</day>
<month>06</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>13</volume>
<elocation-id>1600945</elocation-id>
<history>
<date date-type="received">
<day>27</day>
<month>03</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>27</day>
<month>05</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2025 Maliyakkal, Kumar, Bhowmik, Vishwakarma, Yadav and Mathew.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Maliyakkal, Kumar, Bhowmik, Vishwakarma, Yadav and Mathew</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>
<italic>Trypanosoma cruzi</italic> is the cause of Chagas disease (CD), a major health issue that affects 6&#x2013;7 million individuals globally. Once considered a local problem, migration and non-vector transmission have caused it to spread. Efforts to eliminate CD remain challenging due to insufficient awareness, inadequate diagnostic tools, and limited access to healthcare, despite its classification as a neglected tropical disease (NTD) by the WHO. One of the foremost concerns remains the development of safer and more effective anti-Chagas therapies. In our study, we developed a standardized and robust machine learning-driven QSAR (ML-QSAR) model using a dataset of 1,183 <italic>Trypanosoma cruzi</italic> inhibitors curated from the ChEMBL database to speed up the drug discovery process. Following the calculation of molecular descriptors and feature selection approaches, Support Vector Machine (SVM), Artificial Neural Network (ANN), and Random Forest (RF) models were developed and optimized to elucidate and predict the inhibition mechanism of novel inhibitors. The ANN-driven QSAR model utilizing CDK fingerprints exhibited the highest performance, proven by a Pearson correlation coefficient of 0.9874 for the training set and 0.6872 for the test set, demonstrating exceptional prediction accuracy. Twelve possible inhibitors with pIC<sub>50</sub> &#x2265; 5 were further identified through screening of large chemical libraries using the ANN-QSAR model and ADMET-based filtering approaches. Molecular docking studies revealed that F6609-0134 was the best hit molecule. Finally, the stability and high binding affinity of F6609-0134 were further validated by molecular dynamics simulations and free energy analysis, bolstering its continued assessment as a possible treatment option for Chagas disease.</p>
</abstract>
<kwd-group>
<kwd>Chagas disease</kwd>
<kwd>
<italic>Trypanosoma cruzi</italic>
</kwd>
<kwd>quantitative structure activity relationships</kwd>
<kwd>machine learning</kwd>
<kwd>artificial neural network</kwd>
<kwd>virtual screening</kwd>
<kwd>molecular docking</kwd>
<kwd>molecular dynamics</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Medicinal and Pharmaceutical Chemistry</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>Chagas disease (CD), caused by the protozoan parasite <italic>T</italic>rypanosoma <italic>cruzi</italic>, was diagnosed for the first time in humans in 1909 (<xref ref-type="bibr" rid="B1">Abras et al., 2022</xref>; <xref ref-type="bibr" rid="B33">Rassi et al., 2010</xref>; <xref ref-type="bibr" rid="B32">Rassi et al., 2012</xref>). CD has shifted from a regional to a global health concern, spreading beyond vector-borne transmission. <italic>T. cruzi</italic> infection can occur through blood transfusion, organ transplantation, congenital transmission, and laboratory accidents (<xref ref-type="bibr" rid="B6">Carbajal-de-la-Fuente et al., 2022</xref>). Among these, congenital transmission is particularly alarming, affecting both endemic regions (with vectors) and non-endemic areas (without vectors), making it a growing public health threat worldwide (<xref ref-type="bibr" rid="B6">Carbajal-de-la-Fuente et al., 2022</xref>). CD affects 6&#x2013;7 million people globally, mainly in Latin America. With migration, CD has spread from rural to urban areas and beyond endemic regions, posing a growing global health challenge. In 2010, WHO classified CD as a neglected tropical disease (NTD) and later included it in the 2021&#x2013;2030 roadmap for elimination. Managing NTDs remains challenging, especially in non-endemic regions, due to low awareness, limited diagnostic guidelines, and resource redirection during the COVID-19 pandemic (<xref ref-type="bibr" rid="B11">Fraundorfer, 2024</xref>; <xref ref-type="bibr" rid="B19">Kiehl et al., 2023</xref>). These factors, along with healthcare inaccessibility for vulnerable groups, threaten progress toward 2030 targets. Control efforts focus on vector control and screening, as no vaccine exists. The complex immunology and chronic nature of CD hinder vaccine development, making prevention and early detection crucial in combating the disease (<xref ref-type="bibr" rid="B36">Sarabi Asiabar et al., 2024</xref>). The currently approved drugs, Benznidazole and Nifurtimox, are effective but associated with severe toxicity due to their nitro groups, which generate reactive metabolic radicals (<xref ref-type="bibr" rid="B15">Kannigadu and N&#x2019;Da, 2020</xref>; <xref ref-type="bibr" rid="B38">Trovato et al., 2020</xref>). This leads to adverse effects such as mutagenicity, genotoxicity, and carcinogenicity, limiting their long-term therapy (<xref ref-type="bibr" rid="B15">Kannigadu and N&#x2019;Da, 2020</xref>). Given these challenges, the search for safer and more selective therapeutic alternatives has gained at this momentum. Since the 1990s, researchers have focused on sterol 14&#x3b1;-demethylase (CYP51) inhibitors, which is a crucial target in the sterol biosynthesis pathway of parasite (<xref ref-type="bibr" rid="B8">Choi et al., 2014</xref>; <xref ref-type="bibr" rid="B22">Lepesheva et al., 2011</xref>; <xref ref-type="bibr" rid="B29">Patterson and Fairlamb, 2019</xref>). These inhibitors offer greater selectivity and potentially reduced toxicity, making them promising candidates for improved CD treatment. However, clinical development of these compounds remains challenging, requiring further optimization to balance efficacy and safety. CYP51, a key enzyme in sterol biosynthesis, is essential for parasite survival, making it a promising drug target. Azoles like posaconazole and ravuconazole inhibit CYP51 by interacting with its heme iron, offering potential for selective treatment (<xref ref-type="bibr" rid="B21">Lepesheva et al., 2007</xref>; <xref ref-type="bibr" rid="B27">Parker et al., 2014</xref>; <xref ref-type="bibr" rid="B30">Rabelo et al., 2017</xref>). Other targets include cruzipain, pyrophosphate enzymes, and trypanothione reductase, though many inhibitors have shown high toxicity. Despite extensive research, benznidazole and nifurtimox remain the only FDA-approved drugs. Recent studies suggest piperazine analogues of fenarimol as safer alternatives, highlighting the need for novel and less toxic therapies (<xref ref-type="bibr" rid="B24">Mazzeti et al., 2021</xref>; <xref ref-type="bibr" rid="B35">Salomao et al., 2016</xref>; <xref ref-type="bibr" rid="B43">Zobi and Algul, 2025</xref>). The new derivatives with amide, sulfonamide, aromatic, carbamate, and carbonate substituents were evaluated for their ability to inhibit <italic>T. cruzi in vitro</italic> and showed very promising results (<xref ref-type="bibr" rid="B17">Keenan et al., 2013</xref>; <xref ref-type="bibr" rid="B18">Keenan and Chaplin, 2015</xref>).</p>
<p>Despite 2&#xa0;decades of research, no more effective and less toxic therapeutic alternatives have been identified, and existing drug combinations remain under clinical evaluation (<xref ref-type="bibr" rid="B7">Cheesman et al., 2017</xref>; <xref ref-type="bibr" rid="B13">Harrison et al., 2020</xref>). Cheminformatics and molecular modelling offer a valuable approach, providing cost-effective solutions compared to traditional drug discovery methods (<xref ref-type="bibr" rid="B2">Baldi, 2010</xref>; <xref ref-type="bibr" rid="B3">Bayat Mokhtari et al., 2017</xref>; <xref ref-type="bibr" rid="B37">Siddiqui et al., 2025</xref>). QSAR is a statistical approach that correlates molecular descriptors with biological activity, aiding in the prediction of compounds with more effectiveness (<xref ref-type="bibr" rid="B25">Naithani and Guleria, 2024</xref>; <xref ref-type="bibr" rid="B28">Patel et al., 2014</xref>; <xref ref-type="bibr" rid="B41">Winkler, 2002</xref>). In this study, we developed a robust 2-dimensional machine learning QSAR model using a dataset of <italic>T. cruzi inhibitors</italic> from the ChEMBL database (<ext-link ext-link-type="uri" xlink:href="https://www.ebi.ac.uk/chembl/">https://www.ebi.ac.uk/chembl/</ext-link>) to predict biological activity. The model was trained on multiple molecular descriptors to establish a robust structure-activity relationship, enabling accurate activity predictions for new compounds. To further validate potential candidates, we performed virtual screening using molecular docking to assess binding affinity within the target site. The top-ranked compounds were further subjected to molecular dynamics simulations to evaluate their stability and interactions over time, ensuring their potential effectiveness as novel <italic>T. cruzi</italic> inhibitors.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>2 Materials and methods</title>
<sec id="s2-1">
<title>2.1 Data curation</title>
<p>To construct a machine learning-driven quantitative structural activity relationship model (ML-QSAR), we retrieved a dataset of 1,183 <italic>T. cruzi</italic> inhibitors along with their chemical structures as Simplified Molecular Input Line Entry System (SMILES) and biological data as maximum inhibitory concentration (IC<sub>50</sub>) values from the ChEMBL database (<ext-link ext-link-type="uri" xlink:href="https://chembl.gitbook.io/chembl-interface-documentation/web-services">https://chembl.gitbook.io/chembl-interface-documentation/web-services</ext-link>). The data curation was carried out using chembl web resource client Python module. To ensure a normalized scale for the ML-QSAR model as well as to reduce variability in data analysis, the IC<sub>50</sub> values were converted to pIC<sub>50</sub>, i.e., negative logarithm (base 10) of IC<sub>50</sub>.</p>
</sec>
<sec id="s2-2">
<title>2.2 Molecular descriptor calculation and feature selection</title>
<p>We used padelpy (<ext-link ext-link-type="uri" xlink:href="https://github.com/ecrl/padelpy">https://github.com/ecrl/padelpy</ext-link>), a Python wrapper of the PaDEL-descriptor software, to calculate 1,024 CDK fingerprints and 780 atom pair 2D fingerprints (<xref ref-type="bibr" rid="B42">Yap, 2011</xref>) for the retrieved 1,183 inhibitors. Following the descriptor calculation step, we further implemented variance threshold scores and Pearson correlation analysis-based selection (correlation coefficient &#x3e;0.9) to eliminate the constant and highly correlated features, respectively, from both the fingerprint datasets.</p>
</sec>
<sec id="s2-3">
<title>2.3 ML-QSAR model development and evaluation</title>
<p>We used an 80:20 split ratio for generating the training and test datasets for both fingerprints. We implemented Support Vector Machine (SVM), Artificial Neural Network (ANN), and Random Forest (RF) ML algorithms to develop individual QSAR models for each of the fingerprint datasets using the scikit-learn (<ext-link ext-link-type="uri" xlink:href="https://scikit-learn.org/stable/">https://scikit-learn.org/stable/</ext-link>) Python programming library (<xref ref-type="bibr" rid="B4">Breiman, 2001</xref>; <xref ref-type="bibr" rid="B12">Goel et al., 2023</xref>; <xref ref-type="bibr" rid="B39">Utkin, 2019</xref>). For the SVM model, we implemented the radial basis function (RBF) kernel to capture the non-linear relationships between the molecular fingerprints and biological activity. Additionally, the model was optimized for C (regularization) and gamma (kernel coefficient) parameters. In case of the ANN-driven QSAR model, we implemented a feedforward neural network (FNN) with one hidden layer. We also tuned the number of neurons, activation function (ReLU), and optimizer (Adam) for the ANN model. For the RF-driven QSAR model, an ensemble of decision trees was used along with a feature bagging technique to enhance the predictive power of the model. Additionally, we also optimized the number of estimators (trees), the depth of trees, and the minimum samples per split.</p>
<p>Following the development of the initial model, we further performed principal component analysis (PCA) to assess the distribution of compounds and detect potential outliers in both training and test datasets. PCA was applied to transform the high-dimensional descriptor space into principal components, retaining maximum variance in a low-dimensional space. The first two principal components were further visualized using a scatter plot to inspect cluster formation as well as to detect points deviating from the main distribution. Molecules falling outside the main data clusters were detected as outliers and were removed from further modeling.</p>
<p>Following outlier detection and removal, we further trained the model using grid-based hypertuning and cross-validation metrics using SVM, ANN, and RF algorithms. To determine the best-performing models, we computed a diverse set of statistical metrics for each model: root mean squared error (RMSE), mean squared error (MSE), mean absolute error (MAE), Pearson Correlation coefficient, and 10-fold cross-validation metrics. The model with the lowest RMSE, MSE, and MAE values while retaining a high Pearson Correlation Coefficient was selected as the optimal QSAR model for predicting the inhibition mechanism.</p>
</sec>
<sec id="s2-4">
<title>2.4 Feature elucidation of the ML-QSAR model for rational drug-design</title>
<p>To further enhance the interpretability of the ML-QSAR models for both the fingerprints, we further implemented different feature importance analysis techniques like Variance Importance in Projection Analysis (VIP), Correlation Matrix Analysis, and Shapley Additive Explanations (SHAP) (<ext-link ext-link-type="uri" xlink:href="https://shap.readthedocs.io/en/latest/">https://shap.readthedocs.io/en/latest/</ext-link>) analysis (<xref ref-type="bibr" rid="B26">Nohara et al., 2022</xref>). For the VIP plot analysis, we computed the Partial Least Squares Regression (PLSR) method to rank descriptors based on their contribution to the model (<xref ref-type="bibr" rid="B5">Cao et al., 2017</xref>). For Correlation Matrix analysis, a pairwise correlation matrix was generated to identify the positively and negatively correlated molecular fingerprints for both models. For SHAP analysis, we individually computed SHAP values for each fingerprint to interpret how each molecular fingerprint influenced the pIC<sub>50</sub> values across the datasets. Additionally, we also did cluster analysis for both highly active (pIC<sub>50</sub>&#x2265;7) and weak inactive molecules (pIC<sub>50</sub>&#x2264;4) using Tanimoto Coefficient-based similarity analysis (<ext-link ext-link-type="uri" xlink:href="https://github.com/MunibaFaiza/tanimoto_similarities">https://github.com/MunibaFaiza/tanimoto_similarities</ext-link>). This clustering approach helped in identifying shared molecular features among compounds demonstrating strong or weak inhibition activity. To further enhance the interpretability of the clustering approach, we employed a WordCloud (<ext-link ext-link-type="uri" xlink:href="https://pypi.org/project/wordcloud/">https://pypi.org/project/wordcloud/</ext-link>) approach to visualize the most frequently occurring molecular features among the two groups.</p>
</sec>
<sec id="s2-5">
<title>2.5 Machine learning-driven chemical library screening</title>
<p>To further pave the path for the discovery of novel and more effective drug candidates, we implemented a two-dimensional multiplex modeling to screen large chemical libraries. We curated an Antiprotozoal Screening Compound Library of 8,200 molecules from the Life Chemicals database (<ext-link ext-link-type="uri" xlink:href="https://lifechemicals.com/screening-libraries/targeted-and-focused-screening-libraries/antiprotozoal-library">https://lifechemicals.com/screening-libraries/targeted-and-focused-screening-libraries/antiprotozoal-library</ext-link>). We then implemented the first layer of virtual screening approach, i.e., pharmacokinetic and toxicophore analysis of the molecule library through the ChemBioServer 2.0 (<ext-link ext-link-type="uri" xlink:href="https://chembioserver.vi-seem.eu/">https://chembioserver.vi-seem.eu/</ext-link>). We initially screened the molecules through Lipinski&#x2019;s Rule of Five, Veber&#x2019;s Rule, and Ghose&#x2019;s Filter using the ChemBioServer 2.0. Following this, we again used the ChemBioServer 2.0 to utilize the toxiphore analysis approach to screen out the toxic molecules (<xref ref-type="bibr" rid="B16">Karatzas et al., 2020</xref>). The screened molecules were then subjected to the second layer of virtual screening, i.e., activity prediction and screening using our previously developed ML-QSAR for both the molecular feature datasets, i.e., CDK and atom 2D pair fingerprints.</p>
</sec>
<sec id="s2-6">
<title>2.6 Molecular docking</title>
<p>On the aforementioned conformations, molecular docking studies were performed to investigate residue interactions and binding energy scores of lead molecules from chemical library screening. For the current investigation, the drug target, cruzain enzyme from <italic>T. cruzi</italic> (PDB ID: 1ME3) (<xref ref-type="bibr" rid="B14">Huang et al., 2003</xref>) was retrieved from the protein data bank (<ext-link ext-link-type="uri" xlink:href="https://www.rcsb.org/">https://www.rcsb.org/</ext-link>) with a superior resolution of 1.2&#xa0;&#xc5;, bearing a co-crystallized ligand<italic>.</italic> Missing residues were restored with the glide after the existing ligands were removed and hydrogen atoms were added. The co-crystallized ligand has been redocked to the active site of the 1ME3 to ascertain the docking parameters. Molecular docking employed a three-step approach that comprised protein energy reduction using the Protein Preparation Wizard (PPW) tool, optimization, and pre-processing to create protein crystal structures. LigPrep was utilized to create the ligands, guaranteeing accurate assignment of atom types and protonation states at pH 7.4 &#xb1; 1.0. Hydrogen atoms were added, and the structures underwent bond ordering. Then, using the Receptor grid generating tool (<xref ref-type="bibr" rid="B20">Kumar et al., 2024</xref>; <xref ref-type="bibr" rid="B31">Rampogu et al., 2018</xref>), a grid was created at the binding pocket coordinates (x, y, z) and aligned with a co-crystallized ligand.</p>
</sec>
<sec id="s2-7">
<title>2.7 Molecular dynamics simulation (MDS)</title>
<p>The &#x201c;Desmond V 7.2 package&#x201d; (Schrodinger 2022-4) has been used to conduct MDS to examine how the solvent system affects the structure of the protein-ligand complex. The simulations were done on a Dell Inc. Precision 7,820 Tower running Ubuntu 22.04.1 LTS 64-bit and outfitted with an Intel Xeon (R) Silver 4210R processor and an NVIDIA Corporation GP104GL (RTX A 4000) graphics processing unit. The docked complex&#x2019;s MDS (F6609-0134-1ME3) was performed using the OPLS4 force field. For MDS, the complex is positioned in the middle of an orthorhombic cubic box. After adding SPC water molecules and buffers, the protein atom and the edge of the box are separated by 10&#xa0;&#xc5; using the NPT ensemble. Together with counterions like Na&#x2b; and Cl-injected to randomly neutralize the system, the boundary condition box volume has also been calculated depending on the complex type. To assess domain correlations, a study of the protein-ligand interaction, root mean square deviation (RMSD), and root mean square fluctuation (RMSF) was conducted over all C&#x3b1; atoms during the 200&#xa0;ns MD simulation (<xref ref-type="bibr" rid="B9">da Costa et al., 2022</xref>; <xref ref-type="bibr" rid="B23">Maliyakkal et al., 2024</xref>).</p>
</sec>
</sec>
<sec sec-type="results|discussion" id="s3">
<title>3 Results and discussion</title>
<sec id="s3-1">
<title>3.1 ML-QSAR model development</title>
<p>Following the removal of constant features and highly correlated features using variance threshold and correlation-based feature elimination approaches, we retained 533 molecular features for the CDK fingerprint Dataset, and 25 molecular features for the atom pair 2D fingerprint dataset. Following this, we followed an 80:20 split for both fingerprint datasets before model development, individually for the two fingerprint datasets. The split resulted in 946 molecules in the training set and 237 molecules in the test set. We further implemented SVM, ANN, and RF-based ML algorithms using the generated training and test datasets to develop individual ML-QSARs for two different molecular features, i.e., CDK and atom 2D pair fingerprint datasets.</p>
<p>For the CDK fingerprint model, the ANN-driven QSAR model was found to show the best statistics and model accuracy for the initial run. The model demonstrated a Pearson correlation coefficient of 0.9874 and 0.6872, RMSE of 0.1511 and 0.7271, MSE of 0.0228 and 0.5287, and MAE of 0.0.086 and 0.5614, for training and test datasets, respectively. Following the initial model run, we further implemented PCA analysis to detect the outliers in the datasets. The PCA-based clustering revealed that 119 molecules from both the training and test datasets deviated significantly from the main cluster and were classified as outliers. These molecules were removed from both the training and test datasets before final model deployment. Following outlier removal, we optimized hyperparameters for all 3&#xa0;ML algorithms using randomized search cross-validation. For the RF model, we set the n_estimators between 100 and 700, the max_depth range between 10 and 30, and the min_sample_split range between 2 and 6. For the SVM model, we optimized the C value between 1 and 100, the gamma value range between 0.01 and 0.0001, the epsilon value between 0.01 and 0.2, and kernel type as &#x201c;rbf,&#x201d; &#x201c;poly,&#x201d; &#x201c;sigmoid.&#x201d; Lastly, for the RF model, we optimized hidden layer size to [(128,64), (256,128,64), (512,256,128)], set activation functions to &#x201c;relu&#x201d; and &#x201c;tanh&#x201d;, and solver to &#x201c;adam&#x201d; and &#x201c;sdg.&#x201d; For the RF model, the learning rates and maximum iterations were kept in the range of 0.001&#x2013;0.1 and 1,000&#x2013;2,000, respectively. Even after hyperparameter optimization, the ANN model demonstrated the best performance for the CDK fingerprints by demonstrating a Pearson correlation coefficient of 0.9845 and 0.7683, RMSE of 0.1674 and 0.6167, MSE of 0.028 and 0.3804, and MAE of 0.0915 and 0.4842, for the training and test datasets, respectively. Additionally, it demonstrated a 10-fold cross-validation of 0.7870. Overall, it was evident that the hypertuned CDK fingerprint-driven ANN-QSAR shows the best accuracy and robustness among all the trained models (<xref ref-type="fig" rid="F1">Figure 1</xref>; <xref ref-type="table" rid="T1">Table 1</xref>).</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>
<bold>(A,B)</bold> represent regression plots for the training and test datasets for the best CDK fingerprint-driven ANN-QSAR model, respectively; <bold>(C,D)</bold> represent regression plots for the training and test datasets for the best atom 2D pair fingerprint-driven RF-QSAR model.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g001.tif"/>
</fig>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Statistical metrics of all the generated CDK fingerprint-driven QSAR models. Best model is in bold.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Algorithm</th>
<th align="center">Train RMSE</th>
<th align="center">Test RMSE</th>
<th align="center">CV RMSE</th>
<th align="center">Train MSE</th>
<th align="center">Test MSE</th>
<th align="center">Train MAE</th>
<th align="center">Test MAE</th>
<th align="center">Train pearson</th>
<th align="center">Test pearson</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td colspan="10" align="left">Initial models</td>
</tr>
<tr>
<td align="left">&#x2003;RF</td>
<td align="center">0.2659</td>
<td align="center">0.5986</td>
<td align="center">0.6908</td>
<td align="center">0.0707</td>
<td align="center">0.3583</td>
<td align="center">0.1943</td>
<td align="center">0.4607</td>
<td align="center">0.9656</td>
<td align="center">0.7630</td>
</tr>
<tr>
<td align="left">&#x2003;SVM</td>
<td align="center">0.4449</td>
<td align="center">0.5887</td>
<td align="center">0.6682</td>
<td align="center">0.1980</td>
<td align="center">0.3465</td>
<td align="center">0.2755</td>
<td align="center">0.4493</td>
<td align="center">0.8845</td>
<td align="center">0.7735</td>
</tr>
<tr>
<td align="left">&#x2003;ANN</td>
<td align="center">0.1512</td>
<td align="center">0.7272</td>
<td align="center">0.8266</td>
<td align="center">0.0229</td>
<td align="center">0.5288</td>
<td align="center">0.0868</td>
<td align="center">0.5614</td>
<td align="center">0.9875</td>
<td align="center">0.6872</td>
</tr>
<tr>
<td colspan="10" align="left">Final models (hypertuned and outliers removed)</td>
</tr>
<tr>
<td align="left">&#x2003;RF</td>
<td align="center">0.3413</td>
<td align="center">0.5692</td>
<td align="center">0.6824</td>
<td align="center">0.1165</td>
<td align="center">0.3240</td>
<td align="center">0.2518</td>
<td align="center">0.4417</td>
<td align="center">0.9441</td>
<td align="center">0.7985</td>
</tr>
<tr>
<td align="left">&#x2003;SVM</td>
<td align="center">0.4827</td>
<td align="center">0.6345</td>
<td align="center">0.7237</td>
<td align="center">0.2330</td>
<td align="center">0.4026</td>
<td align="center">0.3029</td>
<td align="center">0.4835</td>
<td align="center">0.8641</td>
<td align="center">0.7415</td>
</tr>
<tr>
<td align="left">&#x2003;ANN (best model)</td>
<td align="center">
<bold>0.1675</bold>
</td>
<td align="center">
<bold>0.6168</bold>
</td>
<td align="center">
<bold>0.7870</bold>
</td>
<td align="center">
<bold>0.0280</bold>
</td>
<td align="center">
<bold>0.3804</bold>
</td>
<td align="center">
<bold>0.0915</bold>
</td>
<td align="center">
<bold>0.4842</bold>
</td>
<td align="center">
<bold>0.9846</bold>
</td>
<td align="center">
<bold>0.7683</bold>
</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>For the atom 2D fingerprint model, the RF-driven QSAR model was found to show the best statistics and model accuracy for the initial run. The model demonstrated a Pearson correlation coefficient of 0.86088 and 0.73814, RMSE of 0.47631 and 0.62489, MSE of 0.22687 and 0.39084, and MAE of 0.33354 and 0.49899, for the training and test datasets, respectively. Additionally, the 10-fold cross-validation for the initial run of the RF-QSAR model was performed. Following the initial model run, we further implemented PCA analysis to detect the outliers in the datasets. The PCA-based clustering revealed that 118 molecules from both the training and test datasets deviated significantly from the main cluster and were classified as outliers. These compounds were removed from both the training and test datasets before final model deployment. Following outlier removal, we optimized hyperparameters for all 3&#xa0;ML algorithms using randomized search cross-validation. We kept the same hyperparameters for the 3&#xa0;ML algorithms that we implemented previously for the CDK fingerprint-based QSAR models. Even after hyperparameter optimization, the RF model demonstrated the best performance for atom 2D pair fingerprints by demonstrating a Pearson correlation coefficient of 0.79293 and 0.69248, RMSE of 0.5438 and 0.63825, MSE of 0.2957 and 0.4073, and MAE of 0.4111 and 0.5184, for training and test datasets, respectively. Additionally, it demonstrated a 10-fold cross-validation of 0.7161. Since there was a reduction in the statistical robustness of the outlier-driven hypertuned model, we implemented the same hyperparameter optimization on the original dataset without outlier removal. We observed that even though the RF algorithm performed better than the outlier-driven hypertuned RF model, it still demonstrated low accuracy and robustness as compared to the initial RF-QSAR model with default parameters. Thereby, we concluded that the initial RF-QSAR model without outlier analysis demonstrated the best performance for the atom 2D pair fingerprint dataset (<xref ref-type="fig" rid="F1">Figure 1</xref>; <xref ref-type="table" rid="T2">Table 2</xref>).</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Statistical metrics of all the generated atom 2D pair fingerprint-driven QSAR models. Best model is bold.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Algorithm</th>
<th align="center">Train RMSE</th>
<th align="center">Test RMSE</th>
<th align="center">CV RMSE</th>
<th align="center">Train MSE</th>
<th align="center">Test MSE</th>
<th align="center">Train MAE</th>
<th align="center">Test MAE</th>
<th align="center">Train pearson</th>
<th align="center">Test pearson</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td colspan="10" align="left">Model (default parameters)</td>
</tr>
<tr>
<td align="left">&#x2003;RF (best model)</td>
<td align="center">
<bold>0.4763</bold>
</td>
<td align="center">
<bold>0.6249</bold>
</td>
<td align="center">
<bold>0.7472</bold>
</td>
<td align="center">
<bold>0.2269</bold>
</td>
<td align="center">
<bold>0.3905</bold>
</td>
<td align="center">
<bold>0.3335</bold>
</td>
<td align="center">
<bold>0.4990</bold>
</td>
<td align="center">
<bold>0.8609</bold>
</td>
<td align="center">
<bold>0.7381</bold>
</td>
</tr>
<tr>
<td align="left">&#x2003;SVM</td>
<td align="center">0.6141</td>
<td align="center">0.6925</td>
<td align="center">0.7331</td>
<td align="center">0.3771</td>
<td align="center">0.4795</td>
<td align="center">0.4222</td>
<td align="center">0.5514</td>
<td align="center">0.7549</td>
<td align="center">0.6655</td>
</tr>
<tr>
<td align="left">&#x2003;ANN</td>
<td align="center">0.4851</td>
<td align="center">0.7449</td>
<td align="center">0.8031</td>
<td align="center">0.2354</td>
<td align="center">0.5548</td>
<td align="center">0.3392</td>
<td align="center">0.5731</td>
<td align="center">0.8595</td>
<td align="center">0.6319</td>
</tr>
<tr>
<td colspan="10" align="left">Model (hypertuned)</td>
</tr>
<tr>
<td align="left">&#x2003;RF</td>
<td align="center">0.4846</td>
<td align="center">0.6237</td>
<td align="center">0.7415</td>
<td align="center">0.2348</td>
<td align="center">0.3891</td>
<td align="center">0.3446</td>
<td align="center">0.5007</td>
<td align="center">0.8561</td>
<td align="center">0.7397</td>
</tr>
<tr>
<td align="left">&#x2003;SVM</td>
<td align="center">0.7839</td>
<td align="center">0.8239</td>
<td align="center">0.8147</td>
<td align="center">0.6145</td>
<td align="center">0.6788</td>
<td align="center">0.5964</td>
<td align="center">0.6568</td>
<td align="center">0.5404</td>
<td align="center">0.4691</td>
</tr>
<tr>
<td align="left">&#x2003;ANN</td>
<td align="center">0.4639</td>
<td align="center">0.6992</td>
<td align="center">0.8298</td>
<td align="center">0.2152</td>
<td align="center">0.4888</td>
<td align="center">0.3047</td>
<td align="center">0.5444</td>
<td align="center">0.8671</td>
<td align="center">0.6698</td>
</tr>
<tr>
<td colspan="10" align="left">Model (hypertuned and outliers removed)</td>
</tr>
<tr>
<td align="left">&#x2003;RF</td>
<td align="center">0.5438</td>
<td align="center">0.6383</td>
<td align="center">0.7162</td>
<td align="center">0.2958</td>
<td align="center">0.4074</td>
<td align="center">0.4112</td>
<td align="center">0.5185</td>
<td align="center">0.7929</td>
<td align="center">0.6925</td>
</tr>
<tr>
<td align="left">&#x2003;SVM</td>
<td align="center">0.6314</td>
<td align="center">0.6979</td>
<td align="center">0.7286</td>
<td align="center">0.3987</td>
<td align="center">0.4870</td>
<td align="center">0.4554</td>
<td align="center">0.5593</td>
<td align="center">0.6877</td>
<td align="center">0.6008</td>
</tr>
<tr>
<td align="left">&#x2003;ANN</td>
<td align="center">0.5910</td>
<td align="center">0.7095</td>
<td align="center">0.7697</td>
<td align="center">0.3493</td>
<td align="center">0.5034</td>
<td align="center">0.4489</td>
<td align="center">0.5630</td>
<td align="center">0.7479</td>
<td align="center">0.6044</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-2">
<title>3.2 Feature elucidation of the ML-QSAR model for rational drug-design</title>
<p>To further interpret the top 20 significant features ML-QSAR model for rational design of novel and more efficient inhibitors, we implemented VIP plot, correlation matrix, and SHAP analysis. For the CDK fingerprint-driven ANN QSAR, we observed that fingerprints, like FP101, FP64, FP480, FP980, and FP35, demonstrated high variance threshold scores in VIP plot analysis, suggesting that the presence of these features in the inhibitor molecule might lead to an increase in biological activity (pIC<sub>50</sub> value). We also observed that fingerprints, FP84, FP109, and FP645, showed low variance threshold scores in the VIP plot, thereby suggesting they might have a negative relationship with biological activity. However, to build a suggestive narrative as to whether these fingerprints in the VIP plot are positively or negatively correlated, we further investigated them through a Pearson correlation matrix and SHAP analysis. Through Pearson correlation matrix analysis, we observed that FP480, FP980, and FP35 demonstrated positive correlation scores towards biological activity, whereas FP84, FP109, and FP645 demonstrated negative correlation scores towards biological activity. Furthermore, from SHAP analysis, it was evident that fingerprints FP480, FP980, and FP64 demonstrated high feature value, thereby suggesting that their presence would lead to an increased pIC<sub>50</sub> value, whereas on the other hand fingerprints FP84, FP109, and FP645, demonstrated negative SHAP values, suggesting the fact that their presence would lead to a decrease in pIC<sub>50</sub> value. Additionally, we did a molecular feature-driven Tanimoto clustering analysis of the high-activity and low-activity molecules of the QSAR dataset. The cluster analysis of the high activity molecules further validated the presence of fingerprints such as FP480, FP64, FP980, and FP35, which were already visualized by Pearson correlation and SHAP analysis plots as positively correlated features. Furthermore, the cluster analysis of the low activity molecules of the QSAR dataset further validated the presence of fingerprints such as FP84, FP109, and FP645, which were already labelled as negatively correlated by the Pearson correlation matrix and SHAP analysis (<xref ref-type="fig" rid="F2">Figures 2</xref>, <xref ref-type="fig" rid="F3">3</xref>, <xref ref-type="fig" rid="F6">6</xref>).</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Representation of top 20 features from VIP plot <bold>(A)</bold> and Pearson correlation plot <bold>(B)</bold> for the CDK fingerprint-driven ANN-QSAR model.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g002.tif"/>
</fig>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Representation of top features from SHAP analysis <bold>(A)</bold> and outlier analysis through PCA plot <bold>(B)</bold> for the CDK fingerprint-driven ANN-QSAR model.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g003.tif"/>
</fig>
<p>For the atom 2D pair fingerprint-driven RF-QSAR model, it was evident that features AD2D91 (presence of N-N at topological distance 2), AD2D705 (presence of C-O at topological distance 10), AD2D336 (presence of O-O at topological distance 5), AD2D169 (presence of N-N at topological distance 3), AD2D102 (presence of O-O at topological distance 2), AD2D248 (presence of N-O at topological distance 4), and AD2D13 (presence of N-N at topological distance 1), demonstrated high variance threshold scores through the VIP plot analysis. We also observed that molecular features like, AD2D92 (presence of N-O at topological distance 2), AD2D704 (presence of C-N at topological distance 10), AD2D12 (presence of C-X at topological distance 1), AD2D247 (presence of N-N at topological distance 4), AD2D326 (presence of N-O at topological distance 5), AD2D170 (presence of N-O at topological distance 3), AD2D482 (presence of N-O at topological distance 7), AD2D626 (presence of C-N at topological distance 9), AD2D404 (presence of N-O at topological distance 6), AD2D325 (presence of N-N at topological distance 5), AD2D549 (presence of C-O at topological distance 8), AD2D627 (presence of C-O at topological distance 9), and AD2D403 (presence of N-N at topological distance 6) demonstrated moderate to low variance threshold score through VIP plot analysis. To further investigate the nature of the correlation of the VIP plot-derived features, we conducted a Pearson correlation matrix and SHAP analysis. Through the Pearson correlation matrix analysis, it was evident that fingerprints AD2D91, AD2D169, AD2D704, AD2D626, and AD2D12 demonstrated positive correlation with high biological activity, whereas fingerprints AD2D170, AD2D248, AD2D326, AD2D92, AD2D102, AD2D325, and AD2D705 demonstrated moderately positive correlation with biological activity. Additionally, the Pearson correlation matrix also demonstrated that fingerprints, AD2D549, AD2D403, AD2D13, AD2D627, AD2D336, AD2D404, and AD2D482 showed negative correlation with biological activity. For SHAP analysis, it was observed that AD2D91, AD2D169, AD2D170, and AD2D12 demonstrated higher SHAP values, suggesting their positive impact on biological activity, whereas AD2D13, AD2D248, AD2D336, and AD2D705 showcased negative SHAP values, thereby demonstrating their negative impact on biological activity. We also did a molecular feature-driven Tanimoto clustering analysis of the high-activity and low-activity molecules of the QSAR dataset. We observed that positively correlated features such as AD2D549, AD2D102, AD2D336, AD2D92, AD2D248, AD2D404, AD2D704, AD2D170, AD2D403, AD2D626, AD2D12, and AD2D326 from the Pearson correlation matrix and SHAP analysis were also found to be present in the cluster analysis of highly active molecules of the dataset. On the other hand, molecular features like AD2D480, which were labelled to be negatively correlated to biological activity in both SHAP and Pearson correlation matrix analysis, were found to be significantly present in cluster analysis of low activity molecules of the QSAR dataset (<xref ref-type="fig" rid="F4">Figures 4</xref>&#x2013;<xref ref-type="fig" rid="F6">6</xref>).</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Representation of the top 20 features from VIP plot <bold>(A)</bold> and Pearson correlation plot <bold>(B)</bold> for the atom 2D pair fingerprint-driven RF-QSAR model.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g004.tif"/>
</fig>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Representation of top features from SHAP analysis <bold>(A)</bold> and outlier analysis through PCA plot <bold>(B)</bold> for the atom 2D pair fingerprint-driven RF-QSAR model.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g005.tif"/>
</fig>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Representation of top features through Tanimoto similarity-driven cluster analysis of highly active <bold>(A)</bold> and low active molecules <bold>(B)</bold> from the CDK fingerprint dataset; highly active <bold>(C)</bold> and low active <bold>(D)</bold> molecules from the atom 2D pair fingerprint dataset.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g006.tif"/>
</fig>
</sec>
<sec id="s3-3">
<title>3.3 Machine learning-driven chemical library screening</title>
<p>We initially screened the Antiprotozoal Screening Compound Library of 8,200 molecules from the Life Chemicals database using the combination of Lipinski&#x2019;s Rule of Five, Veber&#x2019;s Rule, and Ghose&#x2019;s Filter through the ChemBioServer 2.0. A total of 133 molecules that passed through this filtration step were further subjected to toxiphore analysis. A total of 93 out of 133 molecules were found to pass the toxicophore analysis study. Following this, we further predicted the biological activity (pIC<sub>50</sub> value) of the 93 molecules using our previously developed CDK fingerprint-driven RF-QSAR and atom 2&#xa0;days pair fingerprint-driven ANN-QSAR model. We then calculated the cumulative of the predicted biological activities for each molecule from both models to identify the most promising inhibitor. To streamline the drug discovery process further, we identified 12 molecules with pIC<sub>50</sub> &#x2265; 5 (<xref ref-type="table" rid="T3">Table 3</xref>).</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Top 12 hit molecules bioactivity prediction scores using our previously developed ML-QSAR models.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Name</th>
<th align="center">Predicted_pIC<sub>50</sub>_a2d_fp</th>
<th align="center">Predicted_pIC<sub>50</sub>_cdk_fp</th>
<th align="center">Predicted_pIC<sub>50</sub>_cumulative</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">F2207-0115</td>
<td align="center">4.9917</td>
<td align="center">6.9805</td>
<td align="center">5.9861</td>
</tr>
<tr>
<td align="center">F2207-0102</td>
<td align="center">5.2784</td>
<td align="center">6.6665</td>
<td align="center">5.9725</td>
</tr>
<tr>
<td align="center">F6548-1609</td>
<td align="center">5.1616</td>
<td align="center">6.6195</td>
<td align="center">5.8906</td>
</tr>
<tr>
<td align="center">F6548-3996</td>
<td align="center">5.3880</td>
<td align="center">6.1246</td>
<td align="center">5.7563</td>
</tr>
<tr>
<td align="center">F6609-0134</td>
<td align="center">5.9258</td>
<td align="center">5.5255</td>
<td align="center">5.7256</td>
</tr>
<tr>
<td align="center">F6619-3684</td>
<td align="center">5.9071</td>
<td align="center">5.4719</td>
<td align="center">5.6895</td>
</tr>
<tr>
<td align="center">F2014-0155</td>
<td align="center">5.5000</td>
<td align="center">5.8509</td>
<td align="center">5.6754</td>
</tr>
<tr>
<td align="center">F6609-0164</td>
<td align="center">5.5789</td>
<td align="center">5.7068</td>
<td align="center">5.6428</td>
</tr>
<tr>
<td align="center">F3222-1452</td>
<td align="center">4.9646</td>
<td align="center">6.2789</td>
<td align="center">5.6217</td>
</tr>
<tr>
<td align="center">F0507-2033</td>
<td align="center">6.8440</td>
<td align="center">4.313</td>
<td align="center">5.5785</td>
</tr>
<tr>
<td align="center">F1872-0526</td>
<td align="center">4.5384</td>
<td align="center">6.6049</td>
<td align="center">5.5716</td>
</tr>
<tr>
<td align="center">F0676-0414</td>
<td align="center">6.8440</td>
<td align="center">4.2031</td>
<td align="center">5.5236</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-4">
<title>3.4 Molecular docking</title>
<p>Molecular docking studies were conducted to have a better understanding of the lead compound&#x2019;s binding processes. Lead compounds found by virtual screening coupled with the 1ME3 protein, and the docking procedure was confirmed using native ligands. We also did a comparative docking analysis with the native ligand P10 (PubChem CID: 5289091). P10 molecule has already been experimentally tested as an active against <italic>T. cruzi</italic> (BioAssay AID: 977610; BioAssay AID: 1811). In the previous study on the 3D crystal structure of 1ME3, inhibitors were found to form a strong hydrogen bond with His159, part of the canonical catalytic triad (Cys25, His159, Asn175). The P2-position phenylalanine fits into the hydrophobic S2 pocket formed by Leu67, Ala133, and Leu157, with Glu205 rotating to accommodate the side chain. The inhibitor backbone is stabilized by hydrogen bonds with Gly66 and Asp158, along with key water-mediated interactions. Additionally, the nitrogen atom of inhibitors shows potential interaction with the hydroxyl group of Ser61. The compounds mentioned in <xref ref-type="table" rid="T4">Table 4</xref> have docking scores (XP mode) for 1ME3 ranging from &#x2212;3.843 to &#x2212;6.352&#xa0;kcal/mol. With a docking score of &#x2212;6.352&#xa0;kcal/mol, F6609-0134 had the highest binding affinity of all of them, whereas the co-ligand received a value of &#x2212;6.023&#xa0;kcal/mol (<xref ref-type="table" rid="T4">Table 4</xref>). F6609-0134 established a hydrogen connection with Leu157, more precisely with the NH atom of the pyrimidine ring, according to an analysis of the 2-D and 3-D interaction map. Hydrophobic interactions were also noted with Asp158, Gly160, Glu205, Leu67, Met68, Cys25, Trp26, Thr59, Ser61, and Ser64 (<xref ref-type="fig" rid="F7">Figure 7</xref>). Furthermore, the lead chemical demonstrated hydrogen bonding with significant residues, as previously reported in the literature and discussed above regarding the binding pocket (<xref ref-type="bibr" rid="B10">Durrant et al., 2010</xref>; <xref ref-type="bibr" rid="B31">Rampogu et al., 2018</xref>; <xref ref-type="bibr" rid="B34">Rogers et al., 2012</xref>; <xref ref-type="bibr" rid="B40">Wiggers et al., 2013</xref>; <xref ref-type="bibr" rid="B14">Huang et al., 2003</xref>). The discovered hits may be promising lead candidates for the therapy of Chagas disease, as the lead compound had lower binding energies and higher docking scores than the reference compounds.</p>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Docking score of 1ME3 with lead molecules from virtual screening, and native ligand. Best docking score is in bold.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Compound</th>
<th align="center">Docking score (kcal/mol)</th>
<th align="center">Compound</th>
<th align="center">Docking score (kcal/mol)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">F2207-0115</td>
<td align="center">&#x2212;5.052</td>
<td align="center">F6609-0164</td>
<td align="center">&#x2212;4.842</td>
</tr>
<tr>
<td align="center">F2207-0102</td>
<td align="center">&#x2212;4.368</td>
<td align="center">F3222-1452</td>
<td align="center">&#x2212;5.252</td>
</tr>
<tr>
<td align="center">F6548-1609</td>
<td align="center">&#x2212;3.843</td>
<td align="center">F0507-2033</td>
<td align="center">&#x2212;5.389</td>
</tr>
<tr>
<td align="center">F6548-3996</td>
<td align="center">&#x2212;4.236</td>
<td align="center">F1872-0526</td>
<td align="center">&#x2212;4.864</td>
</tr>
<tr>
<td align="center">
<bold>F6609-0134</bold>
</td>
<td align="center">
<bold>&#x2212;6.352</bold>
</td>
<td align="center">F0676-0414</td>
<td align="center">&#x2212;5.479</td>
</tr>
<tr>
<td align="center">F6619-3684</td>
<td align="center">&#x2212;3.868</td>
<td align="center">Co-ligand</td>
<td align="center">&#x2212;6.023</td>
</tr>
<tr>
<td align="center">F2014-0155</td>
<td align="center">&#x2212;4.94</td>
<td align="left"/>
<td align="left"/>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>2-D and 3D interaction of Lead compound F6609-0134 with binding pocket of 1ME3.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g007.tif"/>
</fig>
</sec>
<sec id="s3-5">
<title>3.5 Molecular dynamics</title>
<p>To investigate the flexibility and stability of the docked complex of F6609-0134 at the binding site of the 1ME3 protein in biological situations, MD simulations were performed. We also performed a comparative MD analysis of our hit molecule against P10 molecules (co-crystallized native ligand). MD trajectories were used to calculate protein-ligand interactions as well as RMSD and RMSF. <xref ref-type="fig" rid="F8">Figure 8</xref> shows a number of analyses of the MD trajectory data for the F6609-0134-1ME3 complex.</p>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>Analysis of the inhibitor-ligand complex using MD simulation: RMSD plot (co-crystallized ligand RMSD is shown in orange, and RMSD of F6609-0134 is shown in green); RMSF plot (co-crystallized ligand RMSF is shown in orange, and RMSF of F6609-0134 is shown in green); and analysis of protein-ligand contacts of the MD trajectory of the F6609-0134-1ME3 complex.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g008.tif"/>
</fig>
<sec id="s3-5-1">
<title>3.5.1 Root mean square deviation</title>
<p>According to RMSD <xref ref-type="fig" rid="F8">Figure 8</xref>, the C&#x3b1; atoms of the protein in connection with F6609-0134 and co-ligand had RMSD values ranging from 0.83 to 1.80&#xa0;&#xc5; and 0.74 to 1.52&#xa0;&#xc5;, respectively. This suggests that the ligand-protein complex remained stable throughout the simulation. Except for a slight variation observed between 120 and 130&#xa0;ns, the protein&#x2019;s RMSD remained constant throughout the simulation, and the co-ligands did not differ all that much. Based on the complex&#x2019;s predicted trajectory, the RMSD values of its C&#x3b1; atoms demonstrated the stability of the protein-ligand complex in a dynamic environment. A higher RMSD value indicates unfolding for protein C&#x3b1; atoms, whereas a smaller value indicates compactness. The modest change in the backbone RMSD further supported the equilibration of the protein-ligand combination. The difference between the highest and lowest RMSD values represented the backbone deviation. In summary, the total RMSD of the F6609-0134-1ME3 complex remains consistent and dependable in a fluctuating environment.</p>
</sec>
<sec id="s3-5-2">
<title>3.5.2 Root mean square fluctuation</title>
<p>The flexibility of the protein system was measured during the simulation using the RMSF of each amino acid residue. The RMSF plot showed that differences in N-terminal residues were more noticeable. During the simulation, it was discovered that the co-ligand and compound F6609-0134 interacted with amino acids 19 and 18, respectively, of 1ME3. With a few exceptions (<xref ref-type="fig" rid="F8">Figure 8</xref>; <xref ref-type="table" rid="T5">Table 5</xref>), all of these interacting residues had RMSF values smaller than 1&#xa0;&#xc5;. Certain amino acid residues in the protein-ligand complex are essential for the stability of dynamic processes. The RMSF parameter, which is derived from the MD simulation trajectories, measures the deviation of individual amino acids from the reference or native structure. The RMSF visualization facilitates comprehension of the remaining vibrations in the F6609-0134-1ME3 complex. This finding suggests a solid binding of the lead medication with minor conformational changes within the binding pocket of the target protein, since the main chain and active site residues only slightly varied.</p>
<table-wrap id="T5" position="float">
<label>TABLE 5</label>
<caption>
<p>Amino acid contacts with the ligand and their RMSF value.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">Compound</th>
<th align="center">Amino acids that come into contact with ligands and their RMSF (&#xc5;)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">F6609-0134</td>
<td align="center">Trp26 (0.44&#xa0;&#xc5;), Asp60 (2.02&#xa0;&#xc5;), Ser64 (0.89&#xa0;&#xc5;), Gly65 (0.89&#xa0;&#xc5;), Gly66 (0.63&#xa0;&#xc5;), Leu67 (0.57&#xa0;&#xc5;), Met68 (0.48&#xa0;&#xc5;), Asn69 (0.50&#xa0;&#xc5;), Asn70 (0.57&#xa0;&#xc5;), Ala133 (0.43&#xa0;&#xc5;), Ser155 (1.61&#xa0;&#xc5;), Glu156 (1.18&#xa0;&#xc5;), Gln156 (0.76&#xa0;&#xc5;), Leu157 (0.62&#xa0;&#xc5;), Asp158 (0.9&#xa0;&#xc5;), Leu201 (0.48&#xa0;&#xc5;), Glu204 (0.56&#xa0;&#xc5;), and Glu205 (0.48&#xa0;&#xc5;)</td>
</tr>
<tr>
<td align="center">Co-ligand</td>
<td align="center">Gln19 (0.46&#xa0;&#xc5;), Cys25 (0.37&#xa0;&#xc5;), Trp26 (0.39&#xa0;&#xc5;), Thr59 (0.81&#xa0;&#xc5;), Asp60 (0.55&#xa0;&#xc5;), Ser61 (0.65&#xa0;&#xc5;), Cys63 (0.47&#xa0;&#xc5;), Ser64 (0.54&#xa0;&#xc5;), Gly65 (0.48&#xa0;&#xc5;), Gly66 (0.42&#xa0;&#xc5;), Leu67 (0.42&#xa0;&#xc5;), Met68 (0.38&#xa0;&#xc5;), Asn70 (0.44&#xa0;&#xc5;), Ala133 (0.39&#xa0;&#xc5;), Ala136 (0.46&#xa0;&#xc5;), Leu157 (0.49&#xa0;&#xc5;), Asp158 (0.55&#xa0;&#xc5;), His159 (0.36&#xa0;&#xc5;), and Trp177 (0.49&#xa0;&#xc5;)</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-5-3">
<title>3.5.3 Protein ligand contact analysis</title>
<p>The most prevalent contact types, as determined by MD simulations, were hydrophobic, hydrogen bonding, and polar (water-mediated hydrogen bonding). According to a protein-ligand contact research, Leu67, Ala133, Glu156, Leu157, Asp158, and Glu205 strongly contacted F6609-0134. The simulation results show that compound F6609-0134, Leu167, Gln156, Leu157, Asp157, and Glu205 can stabilize 1ME3 protein because the particular contact is sustained for over 40% of the simulation time (<xref ref-type="fig" rid="F8">Figure 8</xref>). Comparing the ligand&#x2019;s 2-D interaction during docking (<xref ref-type="fig" rid="F9">Figure 9</xref>) with the subsequent simulation reveals similar interactions. The <xref ref-type="fig" rid="F7">Figure 7</xref> simulation result for compound F6609-0134 indicates that it forms a hydrogen bond with amino acid Leu157, which may indicate that it can stabilize the binding pocket of 1ME3.</p>
<fig id="F9" position="float">
<label>FIGURE 9</label>
<caption>
<p>Secondary structure element (SSE) distribution plotted against residue index for 1ME3. The SSE composition across each trajectory frame throughout the simulation for 1ME3.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g009.tif"/>
</fig>
</sec>
<sec id="s3-5-4">
<title>3.5.4 Protein secondary structure elements</title>
<p>The comparative analysis of secondary structure elements (SSE) in the 1ME3 protein highlights notable differences in structural organization and stability over a 200&#xa0;ns molecular dynamics (MD) simulation. The SSE histogram plots reveal that 1ME3 consistently exhibits prominent &#x3b1;-helices (in red) and &#x3b2;-strands (in blue) across its residue indices. This continuous and broader distribution of secondary structures reflects a well-organized and stable protein conformation. As shown in <xref ref-type="fig" rid="F9">Figure 9</xref>, 1ME3 maintains an overall SSE content of 40.54%, comprising 19.45% &#x3b1;-helices and 21.09% &#x3b2;-strands. These values indicate a slightly more ordered and stable secondary structure. Furthermore, the SSE timeline plots (<xref ref-type="fig" rid="F9">Figure 9</xref>) confirm that 1ME3 preserves its secondary structure throughout the simulation, with minimal structural deviations. Overall, the data suggest that the 1ME3 protein retains its structural integrity and demonstrates marginally enhanced conformational stability, as evidenced by higher SSE content and reduced fluctuations during the MD simulation.</p>
</sec>
<sec id="s3-5-5">
<title>3.5.5 Principal component analysis (PCA)</title>
<p>Throughout the simulation, the PCA method was used to examine the protein&#x2019;s conformational distribution and large-scale collective motions within the protein-ligand complex. Using the Desmond script (trj_essential_dynamics.py), Essential Dynamics (ED) analysis calculated the primary components of C&#x3b1; atoms to anticipate the dynamic behavior of the protein. Except for PC1 and PC2 negative modes, phase-space projection along PC1 showed a consistent conformational distribution. RMSD, RMSF, and PCA values derived from MD simulation trajectories verified the stability of the F6609&#x2013;0134&#x2013;1ME3 complex in dynamic states (<xref ref-type="fig" rid="F10">Figure 10</xref>).</p>
<fig id="F10" position="float">
<label>FIGURE 10</label>
<caption>
<p>PCA of F6609-0134-1ME3 protein-ligand complex.</p>
</caption>
<graphic xlink:href="fchem-13-1600945-g010.tif"/>
</fig>
</sec>
<sec id="s3-5-6">
<title>3.5.6 Molecular mechanics/generalized born surface area (MM-GBSA)</title>
<p>The free binding energy of the ideal molecule, F6609-0134, which exhibited the highest docking score and predicted activity, was analyzed based on its molecular dynamics (MD) simulation frames. Over a 0&#x2013;200&#xa0;ns MD trajectory, the total average binding energies were calculated as follows: &#x394;G Bind (&#x2212;39.84&#xa0;kcal/mol), &#x394;G Bind H-bond (&#x2212;0.67&#xa0;kcal/mol), &#x394;G Bind Lipo (&#x2212;14.81&#xa0;kcal/mol), and &#x394;G Bind vdW (&#x2212;38.96&#xa0;kcal/mol). Analysis of these values, as presented in <xref ref-type="table" rid="T6">Table 6</xref>, indicates that &#x394;G Bind and &#x394;G Bind vdW contributed most significantly to the overall binding energy, emphasizing the role of van der Waals interactions in molecular stability.</p>
<table-wrap id="T6" position="float">
<label>TABLE 6</label>
<caption>
<p>Free binding energies of the molecule F6609-0134 shown through MM-GBSA.&#x2a;</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="center">MD snapshot (ns)</th>
<th align="center">&#x394;G bind</th>
<th align="center">&#x394;G bind H-bond</th>
<th align="center">&#x394;G bind lipo</th>
<th align="center">&#x394;G bind vdW</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">0</td>
<td align="center">&#x2212;42.53</td>
<td align="center">&#x2212;0.44</td>
<td align="center">&#x2212;19.28</td>
<td align="center">&#x2212;41.31</td>
</tr>
<tr>
<td align="center">20</td>
<td align="center">&#x2212;44.91</td>
<td align="center">&#x2212;0.32</td>
<td align="center">&#x2212;12.51</td>
<td align="center">&#x2212;44.77</td>
</tr>
<tr>
<td align="center">40</td>
<td align="center">&#x2212;48.29</td>
<td align="center">&#x2212;0.65</td>
<td align="center">&#x2212;11.28</td>
<td align="center">&#x2212;43.58</td>
</tr>
<tr>
<td align="center">60</td>
<td align="center">&#x2212;36.16</td>
<td align="center">&#x2212;0.09</td>
<td align="center">&#x2212;14.18</td>
<td align="center">&#x2212;43.65</td>
</tr>
<tr>
<td align="center">80</td>
<td align="center">&#x2212;43.75</td>
<td align="center">&#x2212;0.28</td>
<td align="center">&#x2212;16.17</td>
<td align="center">&#x2212;44.75</td>
</tr>
<tr>
<td align="center">100</td>
<td align="center">&#x2212;31.46</td>
<td align="center">&#x2212;1.30</td>
<td align="center">&#x2212;12.02</td>
<td align="center">&#x2212;27.83</td>
</tr>
<tr>
<td align="center">120</td>
<td align="center">&#x2212;38.84</td>
<td align="center">&#x2212;1.05</td>
<td align="center">&#x2212;14.46</td>
<td align="center">&#x2212;34.12</td>
</tr>
<tr>
<td align="center">140</td>
<td align="center">&#x2212;40.92</td>
<td align="center">&#x2212;0.10</td>
<td align="center">&#x2212;19.65</td>
<td align="center">&#x2212;41.10</td>
</tr>
<tr>
<td align="center">160</td>
<td align="center">&#x2212;42.44</td>
<td align="center">&#x2212;0.86</td>
<td align="center">&#x2212;17.16</td>
<td align="center">&#x2212;39.31</td>
</tr>
<tr>
<td align="center">180</td>
<td align="center">&#x2212;35.11</td>
<td align="center">&#x2212;1.33</td>
<td align="center">&#x2212;13.02</td>
<td align="center">&#x2212;37.02</td>
</tr>
<tr>
<td align="center">200</td>
<td align="center">&#x2212;33.87</td>
<td align="center">&#x2212;1.00</td>
<td align="center">&#x2212;13.20</td>
<td align="center">&#x2212;31.12</td>
</tr>
<tr>
<td align="center">Average</td>
<td align="center">
<bold>&#x2212;39.84</bold>
</td>
<td align="center">
<bold>&#x2212;0.67</bold>
</td>
<td align="center">
<bold>&#x2212;14.81</bold>
</td>
<td align="center">
<bold>&#x2212;38.96</bold>
</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>&#x2a;kcal/mol.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>Stable van der Waals contacts with important amino acid residues were found by analyzing the &#x394;G Bind vdW values for F6609-0134 interactions with the protein complex. It was discovered that the binding energies derived from MM-GBSA computations based on MD simulation trajectories and docking investigations were consistent. Interestingly, the molecule&#x2019;s low free binding energy suggested that it had a high affinity for the receptor. These results support F6609-0134&#x2019;s potential as a promising inhibitor by indicating that it interacts strongly with 1ME3.</p>
</sec>
</sec>
</sec>
<sec sec-type="conclusion" id="s4">
<title>4 Conclusion</title>
<p>Trypanosoma cruzi causes Chagas disease, a neglected tropical illness that continues to pose a serious threat to world health because of the lack of effective treatments and medication resistance. The potential of machine learning-driven QSAR modeling to speed up the drug discovery process for Chagas disease is demonstrated in this work. The model exhibiting the highest predictive accuracy among those assessed was the ANN-based QSAR model utilizing CDK fingerprints. The stability and high binding affinity of F6609-0134 were confirmed using molecular docking studies and molecular dynamics simulations, which further substantiated its selection as a possible lead molecule. According to these findings, F6609-0134 is a promising therapeutic option for <italic>Trypanosoma cruzi</italic> that needs more experimental support. The molecule should be manufactured or purchased commercially to enhance its potential as an anti-Chagas agent, and thorough <italic>in vitro</italic> tests aimed at <italic>T. cruzi</italic> should be used to evaluate its effectiveness. In the end, these follow-up investigations will promote its development as a treatment candidate for Chagas disease by confirming its trypanocidal efficacy and offering crucial information on cytotoxicity, selectivity, and mechanism of action.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The datasets used in this study were curated from ChEMBL (<ext-link ext-link-type="uri" xmlns:xlink="http://www.w3.org/1999/xlink" xlink:href="https://www.ebi.ac.uk/chembl/">https://www.ebi.ac.uk/chembl/</ext-link>), which is an open-source data repository. The code for generating the ML-assisted QSAR model can be found at: <ext-link ext-link-type="uri" xmlns:xlink="http://www.w3.org/1999/xlink" xlink:href="https://github.com/RatulChemoinformatics/QSAR-Models">https://github.com/RatulChemoinformatics/QSAR-Models</ext-link>.</p>
</sec>
<sec sec-type="author-contributions" id="s6">
<title>Author contributions</title>
<p>NM: Writing &#x2013; review and editing, Writing &#x2013; original draft. SK: Methodology, Supervision, Data curation, Conceptualization, Writing &#x2013; review and editing, Visualization, Writing &#x2013; original draft, Validation, Formal Analysis. RB: Writing &#x2013; original draft, Data curation, Methodology, Writing &#x2013; review and editing, Formal Analysis. HV: Writing &#x2013; original draft. PY: Writing &#x2013; original draft. BM: Writing &#x2013; review and editing, Conceptualization, Writing &#x2013; original draft, Supervision.</p>
</sec>
<sec sec-type="funding-information" id="s7">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research and/or publication of this article. The authors extend their appreciation to the Deanship of Research and Graduate Studies at King Khalid University for funding this work through Small Research Project under grant number RGP1/149/46.</p>
</sec>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="s9">
<title>Generative AI statement</title>
<p>The author(s) declare that no Generative AI was used in the creation of this manuscript.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Abras</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Ballart</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Fern&#xe1;ndez-Ar&#xe9;valo</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Pinazo</surname>
<given-names>M.-J.</given-names>
</name>
<name>
<surname>Gasc&#xf3;n</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Mu&#xf1;oz</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Worldwide control and management of Chagas disease in a new era of globalization: a close look at congenital <italic>Trypanosoma cruzi</italic> infection</article-title>. <source>Clin. Microbiol. Rev.</source> <volume>35</volume> (<issue>2</issue>), <fpage>e0015221</fpage>. <pub-id pub-id-type="doi">10.1128/cmr.00152-21</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Baldi</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Computational approaches for drug design and discovery: an overview</article-title>. <source>Syst. Rev. Pharm.</source> <volume>1</volume> (<issue>1</issue>), <fpage>99</fpage>. <pub-id pub-id-type="doi">10.4103/0975-8453.59519</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bayat Mokhtari</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Homayouni</surname>
<given-names>T. S.</given-names>
</name>
<name>
<surname>Baluch</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Morgatskaya</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Kumar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Das</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group> (<year>2017</year>). <article-title>Combination therapy in combating cancer</article-title>. <source>Oncotarget</source> <volume>8</volume> (<issue>23</issue>), <fpage>38022</fpage>&#x2013;<lpage>38043</lpage>. <pub-id pub-id-type="doi">10.18632/oncotarget.16723</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Breiman</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2001</year>). <article-title>Random forests</article-title>. <source>Mach. Learn.</source> <volume>45</volume> (<issue>1</issue>), <fpage>5</fpage>&#x2013;<lpage>32</lpage>. <pub-id pub-id-type="doi">10.1023/A:1010933404324</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cao</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Deng</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Ensemble partial least squares regression for descriptor selection, outlier detection, applicability domain assessment, and ensemble modeling in QSAR/QSPR modeling</article-title>. <source>J. Chemom.</source> <volume>31</volume> (<issue>11</issue>). <pub-id pub-id-type="doi">10.1002/cem.2922</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Carbajal-de-la-Fuente</surname>
<given-names>A. L.</given-names>
</name>
<name>
<surname>S&#xe1;nchez-Casaccia</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Piccinali</surname>
<given-names>R. V.</given-names>
</name>
<name>
<surname>Provecho</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Salv&#xe1;</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Meli</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Urban vectors of Chagas disease in the American continent: a systematic review of epidemiological surveys</article-title>. <source>PLoS Neglected Trop. Dis.</source> <volume>16</volume> (<issue>12</issue>), <fpage>e0011003</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pntd.0011003</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheesman</surname>
<given-names>M. J.</given-names>
</name>
<name>
<surname>Ilanko</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Blonk</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Cock</surname>
<given-names>I. E.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>Developing new antimicrobial therapies: are synergistic combinations of plant extracts/compounds with conventional antibiotics the solution?</article-title> <source>Pharmacogn. Rev.</source> <volume>11</volume> (<issue>22</issue>), <fpage>57</fpage>&#x2013;<lpage>72</lpage>. <pub-id pub-id-type="doi">10.4103/phrev.phrev_21_17</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Choi</surname>
<given-names>J. Y.</given-names>
</name>
<name>
<surname>Podust</surname>
<given-names>L. M.</given-names>
</name>
<name>
<surname>Roush</surname>
<given-names>W. R.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Drug strategies targeting CYP51 in neglected tropical diseases</article-title>. <source>Chem. Rev.</source> <volume>114</volume> (<issue>22</issue>), <fpage>11242</fpage>&#x2013;<lpage>11271</lpage>. <pub-id pub-id-type="doi">10.1021/cr5003134</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>da Costa</surname>
<given-names>A. P. L.</given-names>
</name>
<name>
<surname>Silva</surname>
<given-names>J. R. A.</given-names>
</name>
<name>
<surname>de Molfetta</surname>
<given-names>F. A.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Computational discovery of sulfonamide derivatives as potential inhibitors of the cruzain enzyme from <italic>T. cruzi</italic> by molecular docking, molecular dynamics and MM/GBSA approaches</article-title>. <source>Mol. Simul.</source> <volume>48</volume> (<issue>18</issue>), <fpage>1678</fpage>&#x2013;<lpage>1687</lpage>. <pub-id pub-id-type="doi">10.1080/08927022.2022.2120625</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Durrant</surname>
<given-names>J. D.</given-names>
</name>
<name>
<surname>Ker&#xe4;nen</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Wilson</surname>
<given-names>B. A.</given-names>
</name>
<name>
<surname>McCammon</surname>
<given-names>J. A.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Computational identification of uncharacterized cruzain binding sites</article-title>. <source>PLoS Neglected Trop. Dis.</source> <volume>4</volume> (<issue>5</issue>), <fpage>e676</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pntd.0000676</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Fraundorfer</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2024</year>). &#x201c;<article-title>Neglected tropical diseases</article-title>,&#x201d; in <source>South American policy regionalism</source> (<publisher-loc>Brazil</publisher-loc>: <publisher-name>Routledge</publisher-name>), <fpage>223</fpage>&#x2013;<lpage>247</lpage>. <pub-id pub-id-type="doi">10.4324/9781003519577-14</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Goel</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Goel</surname>
<given-names>A. K.</given-names>
</name>
<name>
<surname>Kumar</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>The role of artificial neural network and machine learning in utilizing spatial information</article-title>. <source>Spatial Inf. Res.</source> <volume>31</volume> (<issue>3</issue>), <fpage>275</fpage>&#x2013;<lpage>285</lpage>. <pub-id pub-id-type="doi">10.1007/s41324-022-00494-x</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Harrison</surname>
<given-names>J. R.</given-names>
</name>
<name>
<surname>Sarkar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Hampton</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Riley</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Stojanovski</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Sahlberg</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2020</year>). <article-title>Discovery and optimization of a compound series active against <italic>Trypanosoma cruzi</italic>, the causative agent of Chagas disease</article-title>. <source>J. Med. Chem.</source> <volume>63</volume> (<issue>6</issue>), <fpage>3066</fpage>&#x2013;<lpage>3089</lpage>. <pub-id pub-id-type="doi">10.1021/acs.jmedchem.9b01852</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Brinen</surname>
<given-names>L. S.</given-names>
</name>
<name>
<surname>Ellman</surname>
<given-names>J. A.</given-names>
</name>
</person-group> (<year>2003</year>). <article-title>Crystal structures of reversible ketone-based inhibitors of the cysteine protease cruzain</article-title>. <source>Bioorg. Med. Chem.</source> <volume>11</volume> (<issue>1</issue>), <fpage>21</fpage>&#x2013;<lpage>29</lpage>. <pub-id pub-id-type="doi">10.1016/s0968-0896(02)00427-3</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kannigadu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>N&#x2019;Da</surname>
<given-names>D. D.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Recent advances in the synthesis and development of nitroaromatics as anti-infective drugs</article-title>. <source>Curr. Pharm. Des.</source> <volume>26</volume> (<issue>36</issue>), <fpage>4658</fpage>&#x2013;<lpage>4674</lpage>. <pub-id pub-id-type="doi">10.2174/1381612826666200331091853</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Karatzas</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Zamora</surname>
<given-names>J. E.</given-names>
</name>
<name>
<surname>Athanasiadis</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Dellis</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Cournia</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Spyrou</surname>
<given-names>G. M.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>ChemBioServer 2.0: an advanced web server for filtering, clustering and networking of chemical compounds facilitating both drug discovery and repurposing</article-title>. <source>Bioinforma. Oxf. Engl.</source> <volume>36</volume> (<issue>8</issue>), <fpage>2602</fpage>&#x2013;<lpage>2604</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btz976</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Keenan</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Alexander</surname>
<given-names>P. W.</given-names>
</name>
<name>
<surname>Diao</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Best</surname>
<given-names>W. M.</given-names>
</name>
<name>
<surname>Khong</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Kerfoot</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2013</year>). <article-title>Design, structure-activity relationship and <italic>in vivo</italic> efficacy of piperazine analogues of fenarimol as inhibitors of Trypanosoma cruzi</article-title>. <source>Bioorg. and Med. Chem.</source> <volume>21</volume> (<issue>7</issue>), <fpage>1756</fpage>&#x2013;<lpage>1763</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmc.2013.01.050</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Keenan</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Chaplin</surname>
<given-names>J. H.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>A new era for chagas disease drug discovery?</article-title> <source>Prog. Med. Chem.</source> <volume>54</volume>, <fpage>185</fpage>&#x2013;<lpage>230</lpage>. <pub-id pub-id-type="doi">10.1016/bs.pmch.2014.12.001</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kiehl</surname>
<given-names>W. M.</given-names>
</name>
<name>
<surname>Hodo</surname>
<given-names>C. L.</given-names>
</name>
<name>
<surname>Hamer</surname>
<given-names>G. L.</given-names>
</name>
<name>
<surname>Hamer</surname>
<given-names>S. A.</given-names>
</name>
<name>
<surname>Wilkerson</surname>
<given-names>G. K.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Exclusion of horizontal and vertical transmission as major sources of <italic>Trypanosoma cruzi</italic> infections in a breeding colony of rhesus macaques (Macaca mulatta)</article-title>. <source>Comp. Med.</source> <volume>73</volume> (<issue>3</issue>), <fpage>229</fpage>&#x2013;<lpage>241</lpage>. <pub-id pub-id-type="doi">10.30802/AALAS-CM-23-000005</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kumar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Oh</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Prabhakaran</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Awasti</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Mathew</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Isatin-tethered halogen-containing acylhydrazone derivatives as monoamine oxidase inhibitor with neuroprotective effect</article-title>. <source>Sci. Rep.</source> <volume>14</volume> (<issue>1</issue>), <fpage>1264</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-024-51728-x</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lepesheva</surname>
<given-names>G. I.</given-names>
</name>
<name>
<surname>Ott</surname>
<given-names>R. D.</given-names>
</name>
<name>
<surname>Hargrove</surname>
<given-names>T. Y.</given-names>
</name>
<name>
<surname>Kleshchenko</surname>
<given-names>Y. Y.</given-names>
</name>
<name>
<surname>Schuster</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Nes</surname>
<given-names>W. D.</given-names>
</name>
<etal/>
</person-group> (<year>2007</year>). <article-title>Sterol 14&#x3b1;-demethylase as a potential target for antitrypanosomal therapy: enzyme inhibition and parasite cell growth</article-title>. <source>Chem. Biol.</source> <volume>14</volume> (<issue>11</issue>), <fpage>1283</fpage>&#x2013;<lpage>1293</lpage>. <pub-id pub-id-type="doi">10.1016/j.chembiol.2007.10.011</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lepesheva</surname>
<given-names>G. I.</given-names>
</name>
<name>
<surname>Villalta</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Waterman</surname>
<given-names>M. R.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>Targeting <italic>Trypanosoma cruzi</italic> sterol 14&#x3b1;-demethylase (CYP51)</article-title>. <source>Adv. Parasitol.</source> <volume>75</volume>, <fpage>65</fpage>&#x2013;<lpage>87</lpage>. <pub-id pub-id-type="doi">10.1016/B978-0-12-385863-4.00004-6</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Maliyakkal</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Oh</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Kumar</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Gahori</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Tengli</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Beeran</surname>
<given-names>A. A.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>Synthesis, biochemistry, and <italic>in silico</italic> investigations of isatin-based hydrazone derivatives as monoamine oxidase inhibitors</article-title>. <source>Appl. Biol. Chem.</source> <volume>67</volume> (<issue>1</issue>), <fpage>63</fpage>. <pub-id pub-id-type="doi">10.1186/s13765-024-00917-3</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mazzeti</surname>
<given-names>A. L.</given-names>
</name>
<name>
<surname>Capelari-Oliveira</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Bahia</surname>
<given-names>M. T.</given-names>
</name>
<name>
<surname>Mosqueira</surname>
<given-names>V. C. F.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Review on experimental treatment strategies against <italic>Trypanosoma cruzi</italic>
</article-title>. <source>J. Exp. Pharmacol.</source> <volume>13</volume>, <fpage>409</fpage>&#x2013;<lpage>432</lpage>. <pub-id pub-id-type="doi">10.2147/JEP.S267378</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Naithani</surname>
<given-names>U.</given-names>
</name>
<name>
<surname>Guleria</surname>
<given-names>V.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Integrative computational approaches for discovery and evaluation of lead compound for drug design</article-title>. <source>Front. Drug Discov.</source> <volume>4</volume>. <pub-id pub-id-type="doi">10.3389/fddsv.2024.1362456</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nohara</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Matsumoto</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Soejima</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Nakashima</surname>
<given-names>N.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Explanation of machine learning models using shapley additive explanation and application for real data in hospital</article-title>. <source>Comput. Methods Programs Biomed.</source> <volume>214</volume>, <fpage>106584</fpage>. <pub-id pub-id-type="doi">10.1016/j.cmpb.2021.106584</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Parker</surname>
<given-names>J. E.</given-names>
</name>
<name>
<surname>Warrilow</surname>
<given-names>A. G. S.</given-names>
</name>
<name>
<surname>Price</surname>
<given-names>C. L.</given-names>
</name>
<name>
<surname>Mullins</surname>
<given-names>J. G. L.</given-names>
</name>
<name>
<surname>Kelly</surname>
<given-names>D. E.</given-names>
</name>
<name>
<surname>Kelly</surname>
<given-names>S. L.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>Resistance to antifungals that target CYP51</article-title>. <source>J. Chem. Biol.</source> <volume>7</volume> (<issue>4</issue>), <fpage>143</fpage>&#x2013;<lpage>161</lpage>. <pub-id pub-id-type="doi">10.1007/s12154-014-0121-1</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Patel</surname>
<given-names>H. M.</given-names>
</name>
<name>
<surname>Noolvi</surname>
<given-names>M. N.</given-names>
</name>
<name>
<surname>Sharma</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Jaiswal</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Bansal</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lohan</surname>
<given-names>S.</given-names>
</name>
<etal/>
</person-group> (<year>2014</year>). <article-title>Quantitative structure&#x2013;activity relationship (QSAR) studies as strategic approach in drug discovery</article-title>. <source>Med. Chem. Res.</source> <volume>23</volume> (<issue>12</issue>), <fpage>4991</fpage>&#x2013;<lpage>5007</lpage>. <pub-id pub-id-type="doi">10.1007/s00044-014-1072-3</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Patterson</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Fairlamb</surname>
<given-names>A. H.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Current and future prospects of nitro-compounds as drugs for trypanosomiasis and leishmaniasis</article-title>. <source>Curr. Med. Chem.</source> <volume>26</volume> (<issue>23</issue>), <fpage>4454</fpage>&#x2013;<lpage>4475</lpage>. <pub-id pub-id-type="doi">10.2174/0929867325666180426164352</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rabelo</surname>
<given-names>V. W.</given-names>
</name>
<name>
<surname>Santos</surname>
<given-names>T. F.</given-names>
</name>
<name>
<surname>Terra</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Santana</surname>
<given-names>M. V.</given-names>
</name>
<name>
<surname>Castro</surname>
<given-names>H. C.</given-names>
</name>
<name>
<surname>Rodrigues</surname>
<given-names>C. R.</given-names>
</name>
<etal/>
</person-group> (<year>2017</year>). <article-title>Targeting CYP51 for drug design by the contributions of molecular modeling</article-title>. <source>Fundam. Clin. Pharmacol.</source> <volume>31</volume> (<issue>1</issue>), <fpage>37</fpage>&#x2013;<lpage>53</lpage>. <pub-id pub-id-type="doi">10.1111/fcp.12230</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rampogu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Baek</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Son</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Park</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Zeb</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>Discovery of non-peptidic compounds against chagas disease applying pharmacophore guided molecular modelling approaches</article-title>. <source>Mol. Basel Switz.</source> <volume>23</volume> (<issue>12</issue>), <fpage>3054</fpage>. <pub-id pub-id-type="doi">10.3390/molecules23123054</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rassi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Rassi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Marcondes de Rezende</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>American trypanosomiasis (Chagas disease)</article-title>. <source>Infect. Dis. Clin. N. Am.</source> <volume>26</volume> (<issue>2</issue>), <fpage>275</fpage>&#x2013;<lpage>291</lpage>. <pub-id pub-id-type="doi">10.1016/j.idc.2012.03.002</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rassi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Rassi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Marin-Neto</surname>
<given-names>J. A.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Chagas disease</article-title>. <source>Lancet London Engl.</source> <volume>375</volume> (<issue>9723</issue>), <fpage>1388</fpage>&#x2013;<lpage>1402</lpage>. <pub-id pub-id-type="doi">10.1016/S0140-6736(10)60061-X</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rogers</surname>
<given-names>K. E.</given-names>
</name>
<name>
<surname>Ker&#xe4;nen</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Durrant</surname>
<given-names>J. D.</given-names>
</name>
<name>
<surname>Ratnam</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Doak</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Arkin</surname>
<given-names>M. R.</given-names>
</name>
<etal/>
</person-group> (<year>2012</year>). <article-title>Novel cruzain inhibitors for the treatment of Chagas&#x2019; disease</article-title>. <source>Chem. Biol. Drug Des.</source> <volume>80</volume> (<issue>3</issue>), <fpage>398</fpage>&#x2013;<lpage>405</lpage>. <pub-id pub-id-type="doi">10.1111/j.1747-0285.2012.01416.x</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Salomao</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Menna-Barreto</surname>
<given-names>R. F. S.</given-names>
</name>
<name>
<surname>de Castro</surname>
<given-names>S. L.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Stairway to heaven or hell? Perspectives and limitations of chagas disease chemotherapy</article-title>. <source>Curr. Top. Med. Chem.</source> <volume>16</volume> (<issue>20</issue>), <fpage>2266</fpage>&#x2013;<lpage>2289</lpage>. <pub-id pub-id-type="doi">10.2174/1568026616666160413125049</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sarabi Asiabar</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Jabbari</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Rezapour</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Jabbari Khanbebin</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Atafimanesh</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Mazaheri</surname>
<given-names>E.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>Policy and executive barriers in preventing and eradicating neglected tropical diseases: a systematic review</article-title>. <source>Int. J. Prev. Med.</source> <volume>15</volume>, <fpage>49</fpage>. <pub-id pub-id-type="doi">10.4103/ijpvm.ijpvm_251_23</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Siddiqui</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Yadav</surname>
<given-names>C. S.</given-names>
</name>
<name>
<surname>Akil</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Faiyyaz</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Khan</surname>
<given-names>A. R.</given-names>
</name>
<name>
<surname>Ahmad</surname>
<given-names>N.</given-names>
</name>
<etal/>
</person-group> (<year>2025</year>). <article-title>Artificial intelligence in computer-aided drug design (CADD) tools for the finding of potent biologically active Small molecules: traditional to modern approach</article-title>. <source>Comb. Chem. High Throughput Screen.</source> <volume>28</volume>. <pub-id pub-id-type="doi">10.2174/0113862073334062241015043343</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Trovato</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Sartorius</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>D&#x2019;Apice</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Manco</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>De Berardinis</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Viral emerging diseases: challenges in developing vaccination strategies</article-title>. <source>Front. Immunol.</source> <volume>11</volume>, <fpage>2130</fpage>. <pub-id pub-id-type="doi">10.3389/fimmu.2020.02130</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Utkin</surname>
<given-names>V.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>An imprecise extension of SVM-based machine learning models</article-title>. <source>Neurocomputing</source> <volume>331</volume>, <fpage>18</fpage>&#x2013;<lpage>32</lpage>. <pub-id pub-id-type="doi">10.1016/j.neucom.2018.11.053</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wiggers</surname>
<given-names>H. J.</given-names>
</name>
<name>
<surname>Rocha</surname>
<given-names>J. R.</given-names>
</name>
<name>
<surname>Fernandes</surname>
<given-names>W. B.</given-names>
</name>
<name>
<surname>Sesti-Costa</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Carneiro</surname>
<given-names>Z. A.</given-names>
</name>
<name>
<surname>Cheleski</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2013</year>). <article-title>Non-peptidic cruzain inhibitors with trypanocidal activity discovered by virtual screening and <italic>in vitro</italic> assay</article-title>. <source>PLoS Neglected Trop. Dis.</source> <volume>7</volume> (<issue>8</issue>), <fpage>e2370</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pntd.0002370</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Winkler</surname>
<given-names>D. A.</given-names>
</name>
</person-group> (<year>2002</year>). <article-title>The role of quantitative structure--activity relationships (QSAR) in biomolecular discovery</article-title>. <source>Briefings Bioinforma.</source> <volume>3</volume> (<issue>1</issue>), <fpage>73</fpage>&#x2013;<lpage>86</lpage>. <pub-id pub-id-type="doi">10.1093/bib/3.1.73</pub-id>
</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yap</surname>
<given-names>C. W.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>PaDEL-descriptor: an open source software to calculate molecular descriptors and fingerprints</article-title>. <source>J. Comput. Chem.</source> <volume>32</volume> (<issue>7</issue>), <fpage>1466</fpage>&#x2013;<lpage>1474</lpage>. <pub-id pub-id-type="doi">10.1002/jcc.21707</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zobi</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Algul</surname>
<given-names>O.</given-names>
</name>
</person-group> (<year>2025</year>). <article-title>The significance of mono- and dual-effective agents in the development of new antifungal strategies</article-title>. <source>Chem. Biol. Drug Des.</source> <volume>105</volume> (<issue>1</issue>), <fpage>e70045</fpage>. <pub-id pub-id-type="doi">10.1111/cbdd.70045</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>