<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Chem.</journal-id>
<journal-title>Frontiers in Chemistry</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Chem.</abbrev-journal-title>
<issn pub-type="epub">2296-2646</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1137444</article-id>
<article-id pub-id-type="doi">10.3389/fchem.2023.1137444</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Chemistry</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Combining machine learning and structure-based approaches to develop oncogene PIM kinase inhibitors</article-title>
<alt-title alt-title-type="left-running-head">Almukadi et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fchem.2023.1137444">10.3389/fchem.2023.1137444</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Almukadi</surname>
<given-names>Haifa</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="fn" rid="fn1">
<sup>&#x2020;</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Jadkarim</surname>
<given-names>Gada Ali</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Mohammed</surname>
<given-names>Arif</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Almansouri</surname>
<given-names>Majid</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2164547/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Sultana</surname>
<given-names>Nasreen</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Shaik</surname>
<given-names>Noor Ahmad</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/84160/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Banaganapalli</surname>
<given-names>Babajan</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<xref ref-type="fn" rid="fn1">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/138268/overview"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Department of Pharmacology and Toxicology</institution>, <institution>Faculty of Pharmacy</institution>, <institution>King Abdulaziz University</institution>, <addr-line>Jeddah</addr-line>, <country>Saudi Arabia</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Department of Genetic Medicine</institution>, <institution>Faculty of Medicine</institution>, <institution>King Abdulaziz University</institution>, <addr-line>Jeddah</addr-line>, <country>Saudi Arabia</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Department of Biology</institution>, <institution>College of Science</institution>, <institution>University of Jeddah</institution>, <addr-line>Jeddah</addr-line>, <country>Saudi Arabia</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Department of Clinical Biochemistry</institution>, <institution>Faculty of Medicine</institution>, <institution>King Abdulaziz University</institution>, <addr-line>Jeddah</addr-line>, <country>Saudi Arabia</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Department of Biotechnology</institution>, <institution>Acharya Nagarjuna University</institution>, <addr-line>Guntur</addr-line>, <country>India</country>
</aff>
<aff id="aff6">
<sup>6</sup>
<institution>Princess Al-Jawhara Al-Brahim Center of Excellence in Research of Hereditary Disorders</institution>, <institution>King Abdulaziz University</institution>, <addr-line>Jeddah</addr-line>, <country>Saudi Arabia</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/387261/overview">Khurshid Ahmad</ext-link>, Yeungnam University, Republic of Korea</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2165435/overview">Danishuddin Nan</ext-link>, Independent researcher, Republic of Korea</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/716433/overview">Xiaodong Ma</ext-link>, Anhui University of Chinese Medicine, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/149868/overview">C. George Priya Doss</ext-link>, VIT University, India</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Noor Ahmad Shaik, <email>nshaik@kau.edu.sa</email>; Nasreen Sultana, <email>nasreensci@gmail.com</email>; Babajan Banaganapalli, <email>bbabajan@kau.edu.sa</email>
</corresp>
<fn fn-type="equal" id="fn1">
<label>
<sup>&#x2020;</sup>
</label>
<p>These authors have contributed equally to this work</p>
</fn>
<fn fn-type="other">
<p>This article was submitted to Medicinal and Pharmaceutical Chemistry, a section of the journal Frontiers in Chemistry</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>10</day>
<month>03</month>
<year>2023</year>
</pub-date>
<pub-date pub-type="collection">
<year>2023</year>
</pub-date>
<volume>11</volume>
<elocation-id>1137444</elocation-id>
<history>
<date date-type="received">
<day>04</day>
<month>01</month>
<year>2023</year>
</date>
<date date-type="accepted">
<day>09</day>
<month>02</month>
<year>2023</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2023 Almukadi, Jadkarim, Mohammed, Almansouri, Sultana, Shaik and Banaganapalli.</copyright-statement>
<copyright-year>2023</copyright-year>
<copyright-holder>Almukadi, Jadkarim, Mohammed, Almansouri, Sultana, Shaik and Banaganapalli</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>
<bold>Introduction:</bold> PIM kinases are targets for therapeutic intervention since they are associated with a number of malignancies by boosting cell survival and proliferation. Over the past years, the rate of new PIM inhibitors discovery has increased significantly, however, new generation of potent molecules with the right pharmacologic profiles were in demand that can probably lead to the development of Pim kinase inhibitors that are effective against human cancer.</p>
<p>
<bold>Method:</bold> In the current study, a machine learning and structure based approaches were used to generate novel and effective chemical therapeutics for PIM-1 kinase. Four different machine learning methods, namely, support vector machine, random forest, k-nearest neighbour and XGBoost have been used for the development of models. Total, 54 Descriptors have been selected using the Boruta method.</p>
<p>
<bold>Results:</bold> SVM, Random Forest and XGBoost shows better performance as compared to k-NN. An ensemble approach was implemented and, finally, four potential molecules (CHEMBL303779, CHEMBL690270, MHC07198, and CHEMBL748285) were found to be effective for the modulation of PIM-1 activity. Molecular docking and molecular dynamic simulation corroborated the potentiality of the selected molecules. The molecular dynamics (MD) simulation study indicated the stability between protein and ligands.</p>
<p>
<bold>Discussion:</bold> Our findings suggest that the selected models are robust and can be potentially useful for facilitating the discovery against PIM kinase.</p>
</abstract>
<kwd-group>
<kwd>PIM kinase</kwd>
<kwd>classification models</kwd>
<kwd>virtual screening</kwd>
<kwd>molecular docking</kwd>
<kwd>cancer drug treatment</kwd>
</kwd-group>
</article-meta>
</front>
<body>
<sec id="s1">
<title>Introduction</title>
<p>Proto-oncogene PIM-1 kinase is a member of the serine/threonine protein kinase family (<xref ref-type="bibr" rid="B29">Narlik-Grassow et al., 2014</xref>). PIM kinases are involved in cancer cell survival, proliferation, and tumor growth and are overexpressed in a number of hematological malignancies, in addition to solid cancers such as pancreatic, prostate, and colon cancers (<xref ref-type="bibr" rid="B3">Amson et al., 1989</xref>; <xref ref-type="bibr" rid="B25">Li et al., 2006</xref>; <xref ref-type="bibr" rid="B30">Nawijn et al., 2011</xref>). PIM-1, PIM-2, and PIM-3 are the three highly homologous genes that make up the PIM family. This kinase family is highly homologous with the kinase domains, especially in the linker region and the ATP-binding sites (<xref ref-type="bibr" rid="B51">Warfel and Kraft, 2015</xref>). These enzymes are constitutively expressed in tumors and are becoming more widely acknowledged as crucial survival signal mediators in malignancies, stress responses, and neurological development. PIM-1 kinase is a genuine oncogene that is the focus of drug development research initiatives since it has been linked to the emergence of leukemias, lymphomas, and prostate cancer (<xref ref-type="bibr" rid="B23">Li et al., 2011</xref>; <xref ref-type="bibr" rid="B22">Le et al., 2015</xref>; <xref ref-type="bibr" rid="B19">Huang et al., 2022</xref>). PIM kinases regulate the network of signaling pathways that are critical for tumorigenesis and development, making them attractive drug targets (<xref ref-type="bibr" rid="B11">Drygin et al., 2012</xref>; <xref ref-type="bibr" rid="B47">Tursynbay et al., 2016</xref>).</p>
<p>The crystal structure of PIM-1 has been published by numerous independent groups in both the presence and the absence of its inhibitors (<xref ref-type="bibr" rid="B50">Wang et al., 2013</xref>; <xref ref-type="bibr" rid="B31">Nonga et al., 2021</xref>). Structural research on PIM-1 has found a number of distinctive characteristics that set it apart from other kinases with known structures. The catalytic domain of PIM-1 kinase spans amino acid positions 38 to 290 and includes a conserved glycine loop motif at positions 45 to 50, phosphate-binding sites at positions 44 to 52 and 67, and a proton acceptor site at position 167. The hunt for small-molecule ATP-competitive inhibitors with the potential to develop into novel targeted oncology treatments has been sparked by the involvement of the PIM kinases in important cancer hallmarks. The majority of PIM-1 inhibitors have failed to evolve into a new anticancer medication despite having excellent biochemical potency, largely because they were found to have subpar pharmacological qualities (<xref ref-type="bibr" rid="B10">Dakin et al., 2012</xref>; <xref ref-type="bibr" rid="B11">Drygin et al., 2012</xref>; <xref ref-type="bibr" rid="B32">Ogawa et al., 2012</xref>; <xref ref-type="bibr" rid="B48">Vivek et al., 2017</xref>; <xref ref-type="bibr" rid="B55">Zhao et al., 2017</xref>; <xref ref-type="bibr" rid="B33">Park et al., 2021</xref>). Due to their therapeutic value in cancer, the discovery of PIM-1 inhibitors has increasingly attracted much attention in past few years. The rate of new PIM inhibitor discovery has increased significantly, and there has been demand for a new generation of potent molecules with the right pharmacologic profiles that can probably lead to the development of PIM kinase inhibitors that are effective against human cancer.</p>
<p>This work was undertaken to develop machine learning-based classification models to identify a new class of PIM-1 inhibitors. Under this approach, four different machine learning methods were applied to develop the classification models. These models were further used to screen chemical libraries to retrieve novel potent PIM-1 inhibitors. In addition, we also carried out molecular docking and molecular dynamics simulations to investigate the interaction and stability within the catalytic site of PIM-1 kinase. This multistage approach allows us to screen large chemical libraries efficiently and effectively in a reasonable time. Moreover, it can also help us identify novel chemical scaffolds for potent PIM-1 inhibitors.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>Materials and methods</title>
<sec id="s2-1">
<title>Data collection and model building</title>
<p>All chemical compounds with activity against PIM-1 were collected from the literature and the ChEMBL database (<xref ref-type="bibr" rid="B16">Gaulton et al., 2012</xref>). Inorganic and duplicate compounds were removed from the list. Generally, compounds with IC<sub>50</sub> &#x2264; 10&#xa0;&#x3bc;M will likely be &#x201c;active,&#x201d; predicting a large number of active molecules. However, such a high fraction of active compounds cannot be expected from any experimental platform. Therefore, in order to make the most efficient use of costly experimental validation, the optimal model should identify compounds with affinity higher than 10&#xa0;&#x3bc;M. The higher the value, the higher the drug dose needed to achieve the required potency and, thus, the higher the chance of &#x201c;off-target&#x201d; activity. To address this issue, we chose to set the decision boundary at IC<sub>50</sub> &#x2264; 1&#xa0;&#x3bc;M for active molecules. Molecular descriptors were calculated using the PaDEL software (<xref ref-type="bibr" rid="B54">Yap, 2011</xref>). A two-tier selection procedure was applied to select the best descriptors. First, we randomly selected one descriptor from a pair showing &#x3e;0.85 correlation. Second, descriptors were reduced using the Boruta method (<xref ref-type="bibr" rid="B21">Kursa et al., 2010</xref>). We used four different machine learning methods, namely, Support Vector Machine (SVM) (<xref ref-type="bibr" rid="B28">Mitchell, 1997</xref>), random forest (<xref ref-type="bibr" rid="B6">Breiman, 2001</xref>), Extreme Gradient Boosting (XGBoost) (<xref ref-type="bibr" rid="B9">Chen and Guestrin, 2016</xref>), and kappa nearest neighbor (kNN) (<xref ref-type="bibr" rid="B49">Voulgaris and Magoulas, 2008</xref>), to build the classification models. All the classification experiments and calculations were conducted using the R.3.0.2 environment (<ext-link ext-link-type="uri" xlink:href="http://www.R-project.org/">http://www.R-project.org/</ext-link>) and Python (<ext-link ext-link-type="uri" xlink:href="http://www.python.org/">http://www.python.org/</ext-link>) platform. The compounds used in training and test sets are given in <xref ref-type="sec" rid="s10">Supplementary Tables S1 and S2</xref>, respectively.</p>
</sec>
<sec id="s2-2">
<title>Model validation</title>
<p>A receiver operating characteristic (ROC) plot and area under the curve (AUC) were used to assess the performance of the model (<xref ref-type="bibr" rid="B17">Hanley and McNeil, 1983</xref>; <xref ref-type="bibr" rid="B34">Park et al., 2004</xref>). In <xref ref-type="table" rid="T1">Table 1</xref>, the terms precision (Eq. <xref ref-type="disp-formula" rid="e1">1</xref>), recall (Eq. <xref ref-type="disp-formula" rid="e2">2</xref>), accuracy (Eq. <xref ref-type="disp-formula" rid="e3">3</xref>), and F1 score (Eq. <xref ref-type="disp-formula" rid="e4">4</xref>) are defined along with their relationships to the statistical performance calculations used to assess the quality of the model.<disp-formula id="e1">
<mml:math id="m1">
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>n</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>p</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>v</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>p</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>v</mml:mi>
<mml:mi>e</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>N</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>v</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>
<disp-formula id="e2">
<mml:math id="m2">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>l</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>p</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>v</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>p</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>v</mml:mi>
<mml:mi>e</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>N</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>v</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(2)</label>
</disp-formula>
<disp-formula id="e3">
<mml:math id="m3">
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>y</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>T</mml:mi>
<mml:mi>N</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>T</mml:mi>
<mml:mi>N</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mi>P</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(3)</label>
</disp-formula>
<disp-formula id="e4">
<mml:math id="m4">
<mml:mrow>
<mml:mi>F</mml:mi>
<mml:mn>1</mml:mn>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>.</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>n</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>X</mml:mi>
<mml:mtext>&#x2009;</mml:mtext>
<mml:mi>R</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>n</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>R</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>l</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mo>.</mml:mo>
</mml:mrow>
</mml:math>
<label>(4)</label>
</disp-formula>
</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Evaluation metrics for the test set.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Method</th>
<th align="left">Descriptors</th>
<th align="left">Precision</th>
<th align="left">Recall</th>
<th align="left">Accuracy (Q)</th>
<th align="left">F1 score</th>
<th align="left">AUC</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td rowspan="3" align="left">XGBoost</td>
<td align="left">All descriptors</td>
<td align="left">0.82</td>
<td align="left">0.81</td>
<td align="left">0.83</td>
<td align="left">0.97</td>
<td align="left">0.89</td>
</tr>
<tr>
<td align="left">Boruta</td>
<td align="left">0.81</td>
<td align="left">0.79</td>
<td align="left">0.85</td>
<td align="left">0.80</td>
<td align="left">0.88</td>
</tr>
<tr>
<td align="left">MACCS</td>
<td align="left">0.80</td>
<td align="left">0.76</td>
<td align="left">0.81</td>
<td align="left">0.77</td>
<td align="left">0.92</td>
</tr>
<tr>
<td rowspan="3" align="left">Random forest</td>
<td align="left">All descriptors</td>
<td align="left">0.85</td>
<td align="left">0.81</td>
<td align="left">0.86</td>
<td align="left">0.98</td>
<td align="left">0.91</td>
</tr>
<tr>
<td align="left">Boruta</td>
<td align="left">0.86</td>
<td align="left">0.81</td>
<td align="left">0.87</td>
<td align="left">0.83</td>
<td align="left">0.92</td>
</tr>
<tr>
<td align="left">MACCS</td>
<td align="left">0.80</td>
<td align="left">0.76</td>
<td align="left">0.82</td>
<td align="left">0.78</td>
<td align="left">0.90</td>
</tr>
<tr>
<td rowspan="3" align="left">SVM</td>
<td align="left">All descriptors</td>
<td align="left">0.74</td>
<td align="left">0.73</td>
<td align="left">0.78</td>
<td align="left">0.86</td>
<td align="left">0.83</td>
</tr>
<tr>
<td align="left">Boruta</td>
<td align="left">0.75</td>
<td align="left">0.72</td>
<td align="left">0.78</td>
<td align="left">0.71</td>
<td align="left">0.82</td>
</tr>
<tr>
<td align="left">MACCS</td>
<td align="left">0.70</td>
<td align="left">0.73</td>
<td align="left">0.70</td>
<td align="left">0.69</td>
<td align="left">0.82</td>
</tr>
<tr>
<td rowspan="3" align="left">kNN</td>
<td align="left">All descriptors</td>
<td align="left">0.77</td>
<td align="left">0.75</td>
<td align="left">0.80</td>
<td align="left">0.75</td>
<td align="left">0.84</td>
</tr>
<tr>
<td align="left">Boruta</td>
<td align="left">0.72</td>
<td align="left">0.67</td>
<td align="left">0.75</td>
<td align="left">0.68</td>
<td align="left">0.78</td>
</tr>
<tr>
<td align="left">MACCS</td>
<td align="left">0.81</td>
<td align="left">0.76</td>
<td align="left">0.82</td>
<td align="left">0.77</td>
<td align="left">0.82</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s2-3">
<title>Applicability domain</title>
<p>In order to highlight the region of the chemical space that contains the chemicals for which the model is expected to make accurate predictions, a well-validated predictive model needs to have a defined applicability domain (AD) (<xref ref-type="bibr" rid="B40">Rakhimbekova et al., 2020</xref>). Any predictive model must verify its constraints in terms of its structural domain and response space. As a result, determining a model&#x2019;s AD and evaluating the accuracy of its predictions are both challenging tasks. These QSAR models typically use the training set to cover a certain chemical space. The model&#x2019;s predictions are accurate if any query compound falls within this definition of AD. If not, the prediction might not conform to the model&#x2019;s presumptions. Principal component analysis (PCA) (<xref ref-type="bibr" rid="B44">Sushko et al., 2010</xref>) has been employed in our work to define the AD of the compounds used in this study.</p>
</sec>
<sec id="s2-4">
<title>Y-randomization</title>
<p>To test the robustness of the proposed models, y-randomization was applied. This technique involves randomly mixing up the values of the target variable in the training set (<xref ref-type="bibr" rid="B41">R&#xfc;cker et al., 2007</xref>; <xref ref-type="bibr" rid="B26">Lipi&#x144;ski and Szurmak, 2017</xref>). The same parameters used in the initial model are then applied to a new prediction generated with the scrambled data. Every estimate of the model&#x2019;s accuracy was recorded. In total, 50% of the compounds in the training set were resampled and used in a 500-run y-randomization test.</p>
</sec>
<sec id="s2-5">
<title>Similarity calculations</title>
<p>The Tanimoto coefficient (Tc) (Eq. <xref ref-type="disp-formula" rid="e5">5</xref>) was computed using MACCS-166 fingerprints to quantify chemical similarity. The active and inactive chemicals in the training set were compared against false and true positive compounds in systematic pairwise similarity computations.<disp-formula id="e5">
<mml:math id="m5">
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mi>C</mml:mi>
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>B</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:mo>.</mml:mo>
</mml:mrow>
</mml:math>
<label>(5)</label>
</disp-formula>
</p>
</sec>
<sec id="s2-6">
<title>Substructure analyses</title>
<p>Molecular substructures related to PIM activity were analyzed using the distribution of MACSS fingerprints in active and inactive compounds (Eq. <xref ref-type="disp-formula" rid="e6">6</xref>).<disp-formula id="e6">
<mml:math id="m6">
<mml:mrow>
<mml:mi>F</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>q</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>y</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mi>N</mml:mi>
</mml:msubsup>
<mml:mi>F</mml:mi>
<mml:mi>P</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>&#x7c;</mml:mo>
<mml:mn>0</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:mfrac>
<mml:mi>X</mml:mi>
<mml:mn>100</mml:mn>
<mml:mo>.</mml:mo>
</mml:mrow>
</mml:math>
<label>(6)</label>
</disp-formula>
</p>
</sec>
<sec id="s2-7">
<title>Analysis of probability scores</title>
<p>Additionally, the probability scores of the developed classification models were examined. In general, a molecule is defined as inactive if its probability score is lower than 0.5, while a compound with a probability score of 0.5 is considered active (<xref ref-type="bibr" rid="B37">Ponzoni et al., 2019</xref>). The more this score approaches 1, the more confident we are in our prediction. Here, we examined the probability score distributions for TP (true positive), TN (true negative), FP (false positive), and FN (false negative) results.</p>
</sec>
<sec id="s2-8">
<title>Chemical database screening</title>
<p>The developed models were used to screen the hits against PIM-1. The NCI library and Maybridge databases were used for virtual screening. The National Cancer Institute maintains a repository of compounds that have been evaluated as potential anticancer agents. These compounds represent unique structural diversity based on synthetic and natural products. The Maybridge library consists of a highly diverse set of over 53,000 lead-like compounds. Maybridge Hit-to-Lead was designed for medicinal chemistry, allowing SAR development and hit-to-lead optimization. The following filters were used to select the hits: Filter 1: compounds predicted to be active by all the validated models; Filter 2: compounds having a probability score; and Filter 3: compounds falling within the chemical space of the training set. These compounds were further processed for molecular docking, followed by molecular dynamics simulations. Finally, compounds with the best affinity and conformance within the active site were selected and analyzed.</p>
</sec>
<sec id="s2-9">
<title>Molecular docking</title>
<p>Molecular docking was implemented to identify the best physical confirmation of inhibitor binding within the active site of PIM-1 kinase. The PIM kinase enzyme structure was taken from the Protein Data Bank (PDB ID: 5KZI). All of the docking simulations for this work were performed using AutoDock Vina (<xref ref-type="bibr" rid="B45">Trott and OlsonAutoDock, 2009</xref>) with a 1 spacing, default exhaustiveness, and full ligand flexibility. The grid resolution was internally set to 1&#xc5;. We set the number of binding modes to 10 and exhaustiveness to 8. A cubical grid of size 60 &#xd7; 60 &#xd7; 60 size with 0.375&#xa0;&#xc5; spacing was used around the active sites of the protein. To acquire the structure in the PDBQT format, polar hydrogen atoms were added using AutoDock Tools 92.</p>
</sec>
<sec id="s2-10">
<title>Molecular dynamics simulations</title>
<p>Selected best compounds were further subjected to molecular dynamics (MD) simulations using Groningen Machine for Chemical Simulations (GROMACS v5.1.5) (<xref ref-type="bibr" rid="B38">Pronk et al., 2013</xref>). The parameters and coordinate files for PIM-1 kinase and selected potential hit compounds were generated using the CHARMM27 forcefield in GROMACS and PRODRG, respectively. The TIP3P water model was used for each simulation system, which was neutralized by the addition of Na<sup>&#x2b;</sup> ions in a dodecahedron periodic box. Energy minimization was performed for 50,000 nstep using the steepest descent algorithm to avoid steric clashes. Equilibration of each system was performed in two stages: the first phase was carried out with a constant number of particles, volume, and temperature (NVT) ensemble for 500&#xa0;ps at 300&#xa0;K, using the V-rescale thermostat (<xref ref-type="bibr" rid="B7">Bussi et al., 2007</xref>); and in the second phase, the pressure of each system was equilibrated for 500&#xa0;ps at a constant number of particles, pressure, and temperature (NPT) at 1&#xa0;bar using a Parrinello&#x2013;Rahman barostat (<xref ref-type="bibr" rid="B35">Parrinello and Rahman, 1981</xref>). Each equilibrated system was simulated for 30&#xa0;ns under periodic boundary conditions to avoid edge effects. Electrostatic interactions were handled by the particle mesh Ewald (PME) method, while the heavy-atom bonds were restrained using the LINCS algorithm.</p>
</sec>
</sec>
<sec sec-type="results" id="s3">
<title>Results</title>
<sec id="s3-1">
<title>Model development and evaluation</title>
<p>In total, 54 descriptors from the set of 240 were eventually selected using the Boruta method (<xref ref-type="sec" rid="s10">Supplementary Table S3</xref>). All these descriptors belonged to 12 different classes. The descriptors include autocorrelation, information content, atom-type electrotopological state, Burden modified eigenvalues, molecular distance edge, carbon type, and molecular linear free energy relation<italic>.</italic> The models were trained using four machine learning methods (SVM, random forest, XGBoost, and kNN). Evaluation metrics for the developed models are given in <xref ref-type="table" rid="T1">Table 1</xref>, including accuracy, recall, precision, F1 Score (a measure of a model&#x2019;s accuracy, which takes into account both precision and recall), and Area Under the Curve (AUC) values. SVM, random forest, and XGBoost performed than kNN according to these metrics in combination with the selected descriptor set. Among the three, random forest achieved the best accuracy, at 0.87 for the test set (with selected descriptors), as compared to SVM (0.78) and XGBoost (0.84). In addition, these models also had significant AUC values (<xref ref-type="fig" rid="F1">Figure 1</xref>).</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>ROC curves of the models based on four machine learning approaches for <bold>(A)</bold> all descriptors; <bold>(B)</bold> selected descriptors (Boruta method); <bold>(C)</bold> MACCS fingerprints.</p>
</caption>
<graphic xlink:href="fchem-11-1137444-g001.tif"/>
</fig>
</sec>
<sec id="s3-2">
<title>Applicability domain and y-randomization</title>
<p>An applicability domain (AD) analysis was performed to check the reliability of the generated classification models. <xref ref-type="fig" rid="F2">Figure 2</xref> shows a scatter plot of the PC1 and PC2 coordinates derived from the set of selected PIM-1 compound descriptors. The training and test compounds share similar PC1 and PC2 coordinates, suggesting that predictions were within the applicability domain (AD) of both the training and test sets. To check the robustness of the developed models, y-randomization tests were performed (<xref ref-type="bibr" rid="B41">R&#xfc;cker et al., 2007</xref>). Y-randomization test accuracies were found to be lower, and none of the random trials achieved higher scores than our main models (<xref ref-type="fig" rid="F3">Figure 3</xref>). The average accuracy across all randomly generated models were found to be less than 0.58. This confirms that the selected models are robust and reliable and were not generated by chance correlations. A pairwise comparison of the compounds in each cluster was found to reflect reasonable Tanimoto coefficient similarities between them.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Applicability domain plot based on principal component analysis (PCA) for <bold>(A)</bold> training set and <bold>(B)</bold> test set.</p>
</caption>
<graphic xlink:href="fchem-11-1137444-g002.tif"/>
</fig>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Y-randomization models. <bold>(A)</bold> Accuracy; <bold>(B)</bold> AUC values. A total of 500 y-randomization runs were performed.</p>
</caption>
<graphic xlink:href="fchem-11-1137444-g003.tif"/>
</fig>
</sec>
<sec id="s3-3">
<title>Probability analyses</title>
<p>Probability scores of the selected models, reflecting the probability of belonging to each class, were also analyzed. It is known that a compound with a probability score of &#x2265;0.5 is classified as active, whereas a molecule with a probability below &#x3c;0.5 is classified as inactive. As this score approaches 1, the higher the value, the higher the model&#x2019;s confidence in the prediction is (<xref ref-type="bibr" rid="B27">Minerali et al., 2020</xref>; <xref ref-type="bibr" rid="B14">Esposito et al., 2021</xref>). In our study, we analyzed the distribution of probability scores among TN (true negative), FP (false positive), TP (true positive), and FN (false negative) results. For the SVM model, compounds with a probability score of more than 0.80 (an average value) were more likely to be active, whereas compounds with a probability score of 0.36 were more likely to be inactive. In the case of the random forest model, a compound with a probability score of more than 0.87 was more likely to be active, whereas a compound with a probability score of 0.24 was more likely to be inactive. Random forest achieved values of 0.95 and 0.11 for active and inactive compounds, respectively, indicating greater success in predicting compound activity with the desired probability score (<xref ref-type="sec" rid="s10">Supplementary Figure S1</xref>). False positive compounds were predicted with probability scores of 0.63, 0.65, and 0.69 for the random forest, XGBoost, and SVM models, respectively. In contrast, false negative compounds were found to have probability scores of 0.31, 0.42, and 0.14 for the random forest, SVM, and XGBoost models, respectively. Each predictive model&#x2019;s effectiveness in the early recognition of hits was visually evaluated using a cumulative gain plot (<xref ref-type="table" rid="T2">Table 2</xref>). The cumulative gain curve is an evaluation curve that evaluates the model&#x2019;s performance and contrasts the outcomes with a random selection. It displays the percentage of targets identified when taking into account a particular portion of the population that has the highest likelihood of being a target based on the model. The comparison showed that the XGBoost and random forest methods performed better than SVM and kNN in terms of early recognition of hits (<xref ref-type="fig" rid="F4">Figure 4</xref>).</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Probability scores and docking scores of the selected compounds.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="2" align="left">Compound ID</th>
<th colspan="4" align="left">Classifier probability</th>
<th rowspan="2" align="left">Binding energy</th>
</tr>
<tr>
<th align="left">XGBoost</th>
<th align="left">Random forest</th>
<th align="left">SVM</th>
<th align="left">kNN</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">CHEMBL303779</td>
<td align="left">0.82</td>
<td align="left">0.74</td>
<td align="left">0.84</td>
<td align="left">0.78</td>
<td align="left">&#x2212;8.34</td>
</tr>
<tr>
<td align="left">CHEMBL690270</td>
<td align="left">0.76</td>
<td align="left">0.70</td>
<td align="left">0.92</td>
<td align="left">0.85</td>
<td align="left">&#x2212;7.56</td>
</tr>
<tr>
<td align="left">CHEMBL748285</td>
<td align="left">0.72</td>
<td align="left">0.74</td>
<td align="left">0.68</td>
<td align="left">0.63</td>
<td align="left">&#x2212;9.78</td>
</tr>
<tr>
<td align="left">EBM-MPC</td>
<td align="left">0.81</td>
<td align="left">0.75</td>
<td align="left">0.71</td>
<td align="left">0.71</td>
<td align="left">&#x2212;8.45</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Probabilistic distribution plot showing cumulative gain for the developed models.</p>
</caption>
<graphic xlink:href="fchem-11-1137444-g004.tif"/>
</fig>
</sec>
<sec id="s3-4">
<title>MACCS fingerprint analyses</title>
<p>Molecular substructures related to the PIM-1 activity of the compounds can be identified by analyzing the bits in the MACCS fingerprints. We analyzed the MACCS fingerprints showing a reasonable difference between active and inactive compounds (<xref ref-type="sec" rid="s10">Supplementary Table S4</xref>). The occurrence of MACCS fingerprints differed significantly between active and inactive compounds in the training dataset, suggesting that the substructures represented by these features may be closely related to PIM-1 activity. Descriptions and the number of occurrences of these substructures are listed in <xref ref-type="sec" rid="s10">Supplementary Table S4</xref>. It was found that MACCS38, MACCS52, MACCS92, MACCS98, MACCS107, MACCSFP142, <italic>etc.</italic> are prevalent in active molecules. This is consistent with previous studies, which shows that compounds with such functional groups have therapeutic potential against PIM kinase (<xref ref-type="bibr" rid="B46">Tsuganezawa et al., 2012</xref>; <xref ref-type="bibr" rid="B13">El-Hawary et al., 2018</xref>; <xref ref-type="bibr" rid="B33">Park et al., 2021</xref>).</p>
</sec>
<sec id="s3-5">
<title>Database screening and molecular interaction analyses</title>
<p>The NCI and Maybridge databases were used to screen the potential hits from validated models. Commonly predicted active compounds with high probability scores were selected and further filtered out within the applicability domain (AD) of the training set. These compounds were further subjected to molecular docking simulation (<xref ref-type="table" rid="T2">Table 2</xref>). Finally, four compounds (CHEMBL303779, CHEMBL690270, CHEMBL748285, and N-[(1-ethylbenzimidazol-2-yl)methyl]-3-(4-methoxyphenyl)-1H-pyrazole-4-carboxamide (EBM-MPC)) were observed to have reasonable binding affinity and stable interaction with the catalytic residues in the active site (<xref ref-type="table" rid="T3">Table.3</xref> and <xref ref-type="fig" rid="F5">Figure 5</xref>). A literature survey revealed that Leu44, Lys67, Glu121, and Asp186 are crucial for the interaction of inhibitors (<xref ref-type="bibr" rid="B46">Tsuganezawa et al., 2012</xref>; <xref ref-type="bibr" rid="B13">El-Hawary et al., 2018</xref>; <xref ref-type="bibr" rid="B33">Park et al., 2021</xref>). It can be observed in <xref ref-type="fig" rid="F4">Figure 4</xref> that CHEMBL690270, CHEMBL303779, and EBM-MPC form hydrogen bond interactions with Lys67 and hydrophobic interactions with Asp186 (<xref ref-type="fig" rid="F6">Figure 6</xref>). In contrast, CHEMBL748285 forms hydrogen bonds with Asp186 (<xref ref-type="fig" rid="F6">Figure 6</xref>). The quinazoline ring of compounds was involved in multiple p&#x2013;alkyl interactions. In addition, a number of hydrophobic contacts, particularly residues Leu44, Gly47, Phe49, Ile104, and Leu120, stabilize interaction with hits. PIM inhibitors fall into two broad categories: ATP mimetics, which form hydrogen bonds with the glutamate residue that serves as the hinge (Glu121), and non-ATP mimetics, which bind far from the hinge or interact with the hinge through hydrophobic interactions with a number of residues in the specific hydrophobic pocket that serves as the hinge environment (<xref ref-type="bibr" rid="B13">El-Hawary et al., 2018</xref>; <xref ref-type="bibr" rid="B33">Park et al., 2021</xref>). The Tanimoto coefficient (Tc) similarity score of these selected hits was found to be &#x2264; 0.5 with high-activity compounds (<xref ref-type="fig" rid="F5">Figure 5B</xref>).</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Binding mode analysis of the four selected inhibitors.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Compound</th>
<th align="left">Hydrogen bonding</th>
<th align="left">Hydrophobic interaction</th>
<th align="left">H-bond range (&#xc5;)</th>
<th align="left">Hydrophobic interaction range (&#xc5;)</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">CHEMBL303779</td>
<td align="left">Lys67 and Arg122</td>
<td align="left">Gly45, Gly47, Gly48, Phe49, Ala65, Lys67, Ile104, Leu120, Glu121, Arg122, Pro123, Val126, and Leu174</td>
<td align="left">2.7&#x2013;3.2</td>
<td align="left">3.3&#x2013;4.9</td>
</tr>
<tr>
<td align="left">CHEMBL690270</td>
<td align="left">Lys67 and Asp186</td>
<td align="left">Leu44, Gly45, Phe49, Lys67, Ile104, Val126, Asp128, Glu171, Asn172, Leu174, and Asp186</td>
<td align="left">2.4&#x2013;2.6</td>
<td align="left">3.3&#x2013;4.7</td>
</tr>
<tr>
<td align="left">EBM-MPC</td>
<td align="left">Lys67 and Glu121</td>
<td align="left">Gly47, Val52, Lys67, Ile104, Leu120, Glu121, Pro123, Val126, Leu174, and Asp186</td>
<td align="left">2.8&#x2013;3.0</td>
<td align="left">3.6&#x2013;4.89</td>
</tr>
<tr>
<td align="left">CHEMBL748285</td>
<td align="left">Asn172 and Asp186</td>
<td align="left">Leu44, Val52, Phe49, Asn172, Leu174, Leu182, Leu184, and Asp186</td>
<td align="left">1.6&#x2013;3.1</td>
<td align="left">3.6&#x2013;4.4</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Chemical space and similarity analyses for selected compounds. <bold>(A)</bold> Chemical space of selected compounds; <bold>(B)</bold> heat map of the distance matrix for the selected compounds and active compounds in the training set.</p>
</caption>
<graphic xlink:href="fchem-11-1137444-g005.tif"/>
</fig>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Binding mode analyses of selected compounds within the active site of PIM-1 kinase. Active site residues are shown as gray sticks; the protein backbone is shown as a light gray wire; hydrogen bonds are shown with a green dashed line.</p>
</caption>
<graphic xlink:href="fchem-11-1137444-g006.tif"/>
</fig>
</sec>
<sec id="s3-6">
<title>MD simulation analyses</title>
<p>By analyzing 100-ns MD trajectories, the structural changes to PIM-1 upon inhibitor binding were studied. We examined the RMSD of the protein backbone and the RMSF of the protein&#x2019;s alpha-carbon atoms. As shown in <xref ref-type="fig" rid="F6">Figure 6</xref>, all the systems exhibited stability throughout the 100-ns simulation. The average RMSD value for all four systems was observed to be below 0.31&#xa0;nm, which indicated that simulated complexes displayed RMSD values below the threshold. The average RMSD values further showed that the CHEMBL690270 PIM-1 complex displayed less deviation (0.26&#xa0;nm), whereas CHEMBL303779 and CHEMBL748285 demonstrated similar average values of 0.34&#xa0;nm (<xref ref-type="fig" rid="F7">Figure 7A</xref>). RMSF is a significant value, used to characterize each residue&#x2019;s fluctuation rate upon ligand binding. It was observed that the inhibitor binding residues (Leu44, Phe49, Lys67, Glu121, and Asp186) did not fluctuate significantly (<xref ref-type="fig" rid="F7">Figure 7B</xref>).</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Molecular dynamics simulation analyses. <bold>(A)</bold> RMSD plot; <bold>(B)</bold> RMSF plot; <bold>(C)</bold> hydrogen plot for selected compounds to illustrate protein&#x2013;ligand stability during a 100-ns simulation.</p>
</caption>
<graphic xlink:href="fchem-11-1137444-g007.tif"/>
</fig>
</sec>
</sec>
<sec sec-type="discussion" id="s4">
<title>Discussion</title>
<p>This study was designed with the aim of building a classification model to predict potential hits for PIM-1 kinase. Four different machine learning approaches were used to build the models. Our proposed models performed well in terms of accuracy, F1 score, precision, and recall. We used the area under the receiver operating characteristic curve approach to compare classifiers. The ROC curve is a graphical representation that contrasts a classifier&#x2019;s true positive rate and false positive rate at various threshold levels. The area under this curve, or AUC, is thus a useful metric for assessing machine learning algorithms, since it shows the degree of separability (<xref ref-type="bibr" rid="B35">Parrinello and Rahman, 1981</xref>). A ROC curve with a higher AUC value implies greater sensitivity in identifying active molecules and specificity in rejecting inactive compounds (<xref ref-type="fig" rid="F1">Figure 1</xref>). In addition, our study also distinguished and ranked the top 18 variables, including 2D autocorrelation, Burden modified eigenvalues, and topological charge<italic>.</italic> These descriptors have the capacity to distinguish between active and inactive compounds.</p>
<p>QSAR Classification models must undergo an extensive validation process, and the reliability of those models must be objectively determined. The OECD guidelines state that a model must have a clearly defined domain of applicability (<xref ref-type="bibr" rid="B12">Dwyer et al., 2013</xref>). Additionally, the dataset for such models with a defined AD should cover a broad chemical space and a diverse range of structural types. The AD of PIM kinase inhibitors has been defined using a principal component analysis-based approach for model development. A sufficient level of assurance in the produced models can be seen in the 2D plot obtained from the first two PCs, which represents the training and test set compounds, illustrating their structural variety and similar chemical space (<xref ref-type="fig" rid="F2">Figure 2</xref>). To assess the likelihood of a random correlation for a chosen descriptor, y-randomization was utilized. This technique is used to assess the reliability or robustness of QSAR models and is recognized as one of the most effective validation processes (<xref ref-type="bibr" rid="B41">R&#xfc;cker et al., 2007</xref>). By comparing a developed model&#x2019;s performance to the average measure of 500 random models, which are obtained by using the same parameters as those used to construct the original model along with a randomly scrambled target variable class, the statistical significance of the developed model can be examined. The results of the y-randomization tests demonstrated that the models created for this study did not exhibit these connections by chance and that a true structure&#x2013;activity relationship existed (<xref ref-type="fig" rid="F3">Figure 3</xref>).</p>
<p>Fingerprints describe the molecular makeup of a compound. The description of each molecule is given as a string of binary substructures called a fingerprint. The corresponding fingerprint bit is set to 1 if the specified substructure is present in the given molecule; otherwise, it is set to 0. In our study, we used MACCS fingerprints to represent the presence of structures and their representative substructures in active and inactive compounds. These molecules contained MACCS65, MACCS128, and MACCS90. Compounds having such substructures were found to exhibit reasonable levels of activity toward PIM-1 kinase (<xref ref-type="bibr" rid="B2">Aku&#xe9;-G&#xe9;du et al., 2010</xref>; <xref ref-type="bibr" rid="B12">Dwyer et al., 2013</xref>; <xref ref-type="bibr" rid="B18">Hu et al., 2015</xref>; <xref ref-type="bibr" rid="B52">Wurz et al., 2015</xref>; <xref ref-type="bibr" rid="B24">Li et al., 2016</xref>).</p>
<p>To identify potent PIM-1 inhibitors, virtual screening of the NCI and Maybridge databases was performed using the validated models. To gain structural insight relevant to the inhibitory activities of the newly identified inhibitors, their binding modes in the binding site of PIM-1 were examined. <xref ref-type="fig" rid="F6">Figure 6</xref> shows the most stable binding configurations of selected four compounds derived via docking simulations with potent inhibitors. These compounds appear to be accommodated in a similar way in the binding site of PIM1 (<xref ref-type="bibr" rid="B53">Xia et al., 2009</xref>; <xref ref-type="bibr" rid="B1">Abdelaziz et al., 2018</xref>; <xref ref-type="bibr" rid="B20">Ibrahim et al., 2022</xref>). The necessity of the interactions with the hinge region and Gly-loop residues (<xref ref-type="bibr" rid="B39">Qian et al., 2005</xref>; <xref ref-type="bibr" rid="B36">Pogacic et al., 2007</xref>; <xref ref-type="bibr" rid="B46">Tsuganezawa et al., 2012</xref>; <xref ref-type="bibr" rid="B8">Casuscelli et al., 2013</xref>; <xref ref-type="bibr" rid="B15">Fan et al., 2016</xref>; <xref ref-type="bibr" rid="B1">Abdelaziz et al., 2018</xref>; <xref ref-type="bibr" rid="B5">Bima et al., 2022</xref>; <xref ref-type="bibr" rid="B20">Ibrahim et al., 2022</xref>; <xref ref-type="bibr" rid="B43">Shaik et al., 2022</xref>) for tight binding to PIM-1 was also implicated with potent inhibitors (<xref ref-type="bibr" rid="B53">Xia et al., 2009</xref>; <xref ref-type="bibr" rid="B20">Ibrahim et al., 2022</xref>). Moreover, these four compounds can also interact with the activation loop including the Asp186 residue. A hydrophobic cavity is formed among the Ala65, Ile104, Phe187, Val52, Lys67, and Leu120 residues, and this maintains molecular stability through various hydrophobic forces. Similar interactions have also been noted in earlier published investigations, highlighting the significance of these amino acids for the assembly of PIM-1 inhibitor complexes (<xref ref-type="bibr" rid="B46">Tsuganezawa et al., 2012</xref>; <xref ref-type="bibr" rid="B13">El-Hawary et al., 2018</xref>; <xref ref-type="bibr" rid="B33">Park et al., 2021</xref>). Residue Lys67 is known to be significant in stabilizing the interaction with the compound and to play an important role in the catalytic activity of PIM-1 (<xref ref-type="bibr" rid="B36">Pogacic et al., 2007</xref>; <xref ref-type="bibr" rid="B15">Fan et al., 2016</xref>). In our study, we found that all four compounds interacted with Lys67, either with hydrogen bonds or through hydrophobic contact. Compared to the currently available PIM-1 inhibitors, the four selected compounds exhibit low Tanimoto coefficient (Tc) similarities, highlighting their structural novelty and druggability. Moreover, all these compounds were found to have a similar chemical boundary (<xref ref-type="fig" rid="F5">Figure 5</xref>). Therefore, models constructed using these selected descriptors have good interpretability and reliability.</p>
<p>Molecular docking studies were conducted to analyze the binding mode of inhibitors at the PIM-1 catalytic domain. Notably, these inhibitors are positioned in the active site, between the residues Leu44, Gly45, Phe49, Lys67, Ile104, Lys67, Leu172, Leu174, and Asp186 (<xref ref-type="table" rid="T3">Table 3</xref>). These inhibitors were found to have stabilized the complex with hydrogen and hydrophobic interactions with residues, namely, Lys67 and Asp186. This is consistent with earlier research that revealed that these amino acid residues were essential for the catalytic activity of PIM-1 kinase (<xref ref-type="bibr" rid="B39">Qian et al., 2005</xref>; <xref ref-type="bibr" rid="B4">Banaganapalli et al., 2016</xref>; <xref ref-type="bibr" rid="B42">Shaik et al., 2021</xref>; <xref ref-type="bibr" rid="B5">Bima et al., 2022</xref>; <xref ref-type="bibr" rid="B43">Shaik et al., 2022</xref>).</p>
<p>Although molecular docking has strong computational capabilities, its predictions of the shape of the protein&#x2013;ligand binding are frequently inaccurate. Thus, in this study, we performed 100-ns MD simulations to test the stability of the chosen compounds in the PIM-1 binding pocket. It was determined that selected compounds remained stable in the binding pocket, as analyzed through the RMSD, RMSF, and hydrogen bonds. Most notably, stable hydrogen bonds with the residues Lys67 and Asp186 were observed in the complexes with the compounds (namely, CHEMBL748285, and CHEMBL690270).</p>
</sec>
<sec sec-type="conclusion" id="s5">
<title>Conclusion</title>
<p>The PIM kinase family has become a focus of attention in drug discovery. In particular, the search for inhibitors simultaneously targeting PIM-1 isoforms is of great interest because it opens new horizons toward the discovery of new chemicals capable of therapeutically modulating many biochemical pathways involved in the emergence and development of various cancers. In the present study, ensemble learning based on four different machine learning approaches, together with molecular docking and molecular dynamics simulation, was successfully utilized to identify novel scaffold inhibitors against PIM kinase. By combining machine learning and structure-based approaches, it was possible to evaluate the quantitative contributions of the molecules to the activity. This permitted the guided design of four new molecules, predicted to be potential PIM-1 inhibitors. The molecular docking analyses showed that the active inhibitors were able to interact with the amino acids (Lys67, Asp186, Leu44, Glu171, etc.) crucial for catalytic activity of PIM kinase. The interactions were found to be stable, as investigated through 100-ns molecular dynamics simulation.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s6">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/<xref ref-type="sec" rid="s10">Supplementary Material</xref>; further inquiries can be directed to the corresponding authors.</p>
</sec>
<sec id="s7">
<title>Author contributions</title>
<p>HA, NS, and BB: conceptualization; HA, GJ, and BB: data curation; BB and GJ: formal analysis; HA: funding acquisition; HA, GJ, and BB: methodology; HA: project administration; BB, NS, and HA: resources; BB: software; NS: supervision; HA and NS: validation; BB: visualization; HA, BB, AM, MA, NS, and GJ: writing&#x2014;original draft and review.</p>
</sec>
<sec id="s8">
<title>Funding</title>
<p>This project was funded by the Deanship of Scientific Research (DSR) at King Abdulaziz University, under Grant no. G:207-249-1441. The authors therefore acknowledge the DSR for technical and financial support.</p>
</sec>
<sec sec-type="COI-statement" id="s9">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s10">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fchem.2023.1137444/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fchem.2023.1137444/full&#x23;supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Image1.JPEG" id="SM1" mimetype="application/JPEG" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table1.DOC" id="SM2" mimetype="application/DOC" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="DataSheet1.CSV" id="SM3" mimetype="application/CSV" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="DataSheet2.CSV" id="SM4" mimetype="application/CSV" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Abdelaziz</surname>
<given-names>M. E.</given-names>
</name>
<name>
<surname>M El-Miligy</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Fahmy</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Mahran</surname>
<given-names>M. A.</given-names>
</name>
<name>
<surname>A Hazzaa</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Design, synthesis and docking study of pyridine and thieno[2, 3-b] pyridine derivatives as anticancer PIM-1 kinase inhibitors</article-title>. <source>Bioorg Chem.</source> <volume>80</volume>, <fpage>674</fpage>&#x2013;<lpage>692</lpage>. <pub-id pub-id-type="doi">10.1016/j.bioorg.2018.07.024</pub-id>
</citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Aku&#xe9;-G&#xe9;du</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Nauton</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Th&#xe9;ry</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Bain</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Cohen</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Anizon</surname>
<given-names>F.</given-names>
</name>
<etal/>
</person-group> (<year>2010</year>). <article-title>Synthesis, Pim kinase inhibitory potencies and <italic>in vitro</italic> antiproliferative activities of diversely substituted pyrrolo[2, 3-a]carbazoles</article-title>. <source>Bioorg Med. Chem.</source> <volume>18</volume>, <fpage>6865</fpage>&#x2013;<lpage>6873</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmc.2010.07.036</pub-id>
</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Amson</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Sigaux</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Przedborski</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Flandrin</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Givol</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Telerman</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>1989</year>). <article-title>The human protooncogene product p33pim is expressed during fetal hematopoiesis and in diverse leukemias</article-title>. <source>Proc. Natl. Acad. Sci. U. S. A.</source> <volume>86</volume>, <fpage>8857</fpage>&#x2013;<lpage>8861</lpage>.</citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Banaganapalli</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Mohammed</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Khan</surname>
<given-names>I. A.</given-names>
</name>
<name>
<surname>Al-Aama</surname>
<given-names>J. Y.</given-names>
</name>
<name>
<surname>Elango</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Shaik</surname>
<given-names>N. A.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>A computational protein phenotype prediction approach to analyze the deleterious mutations of human MED12 gene</article-title>. <source>J. Cell Biochem.</source> <volume>117</volume>, <fpage>2023</fpage>&#x2013;<lpage>2035</lpage>. <pub-id pub-id-type="doi">10.1002/jcb.25499</pub-id>
</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bima</surname>
<given-names>A. I. H.</given-names>
</name>
<name>
<surname>Elsamanoudy</surname>
<given-names>A. Z.</given-names>
</name>
<name>
<surname>Albaqami</surname>
<given-names>W. F.</given-names>
</name>
<name>
<surname>Khan</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Parambath</surname>
<given-names>S. V.</given-names>
</name>
<name>
<surname>Al-Rayes</surname>
<given-names>N.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Integrative system biology and mathematical modeling of genetic networks identifies shared biomarkers for obesity and diabetes</article-title>. <source>Math. Biosci. Eng.</source> <volume>19</volume>, <fpage>2310</fpage>&#x2013;<lpage>2329</lpage>. <pub-id pub-id-type="doi">10.3934/mbe.2022107</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Breiman</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2001</year>). <article-title>Random forests</article-title>. <source>Mach. Learn.</source> <volume>45</volume>, <fpage>5</fpage>&#x2013;<lpage>32</lpage>. <pub-id pub-id-type="doi">10.1023/A:1010933404324</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bussi</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Donadio</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Parrinell</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>Canonical sampling through velocity rescaling</article-title>. <source>J. Chem. Phys.</source> <volume>126</volume>, <fpage>014101</fpage>. <pub-id pub-id-type="doi">10.1063/1.2408420</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Casuscelli</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Ardini</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Avanzi</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Casale</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Cervi</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>D&#x27;Anello</surname>
<given-names>M.</given-names>
</name>
<etal/>
</person-group> (<year>2013</year>). <article-title>Discovery and optimization of pyrrolo[1, 2-a]pyrazinones leads to novel and selective inhibitors of PIM kinases</article-title>. <source>Bioorg Med. Chem.</source> <volume>21</volume>, <fpage>7364</fpage>&#x2013;<lpage>7380</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmc.2013.09.054</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Guestrin</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>XGBoost: A scalable tree boosting system</article-title>,&#x201d; in <conf-name>Proceedings of the 22Nd ACM SIGKDD International Conference on Knowledge Discovery and Data Mining</conf-name>, <conf-loc>ACM, San Francisc</conf-loc>, <conf-date>March, 2016</conf-date>.</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dakin</surname>
<given-names>L. A.</given-names>
</name>
<name>
<surname>Block</surname>
<given-names>M. H.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Code</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Dowling</surname>
<given-names>J. E.</given-names>
</name>
<name>
<surname>Feng</surname>
<given-names>X.</given-names>
</name>
<etal/>
</person-group> (<year>2012</year>). <article-title>Discovery of novel benzylidene-1, 3-thiazolidine-2, 4-diones as potent and selective inhibitors of the PIM-1, PIM-2, and PIM-3 protein kinases</article-title>. <source>Bioorg Med. Chem. Lett.</source> <volume>22</volume>, <fpage>4599</fpage>&#x2013;<lpage>4604</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmcl.2012.05.098</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Drygin</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Haddach</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Pierre</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Ryckman</surname>
<given-names>D. M.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Potential use of selective and nonselective pim kinase inhibitors for cancer therapy</article-title>. <source>J. Med. Chem.</source> <volume>55</volume>, <fpage>8199</fpage>&#x2013;<lpage>8208</lpage>. <pub-id pub-id-type="doi">10.1021/jm3009234</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dwyer</surname>
<given-names>M. P.</given-names>
</name>
<name>
<surname>Keertika</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Paruch</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Alvarez</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Labroli</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Poker</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2013</year>). <article-title>Discovery of pyrazolo[1, 5-a]pyrimidine-based pim inhibitors: A template-based approach</article-title>. <source>Bioorg Med. Chem. Lett.</source> <volume>23</volume>, <fpage>6178</fpage>&#x2013;<lpage>6182</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmcl.2013.08.110</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>El-Hawary</surname>
<given-names>S. S.</given-names>
</name>
<name>
<surname>Sayed</surname>
<given-names>A. M.</given-names>
</name>
<name>
<surname>Mohammed</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Khanfar</surname>
<given-names>M. A.</given-names>
</name>
<name>
<surname>Rateb</surname>
<given-names>M. E.</given-names>
</name>
<name>
<surname>Mohammed</surname>
<given-names>T. A.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>New pim-1 kinase inhibitor from the Co-culture of two sponge-associated actinomycetes</article-title>. <source>Front. Chem.</source> <volume>6</volume>, <fpage>538</fpage>. <pub-id pub-id-type="doi">10.3389/fchem.2018.00538</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Esposito</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Landrum</surname>
<given-names>G. A.</given-names>
</name>
<name>
<surname>Schneider</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Stiefl</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Riniker</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>GHOST: Adjusting the decision threshold to handle imbalanced data in machine learning</article-title>. <source>J. Chem. Inf. Model</source> <volume>61</volume>, <fpage>2623</fpage>&#x2013;<lpage>2640</lpage>. <pub-id pub-id-type="doi">10.1021/acs.jcim.1c00160</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Fan</surname>
<given-names>Y. B.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>S. Y.</given-names>
</name>
<etal/>
</person-group> (<year>2016</year>). <article-title>Design and synthesis of substituted pyrido[3, 2-d]-1, 2, 3-triazines as potential Pim-1 inhibitors</article-title>. <source>Bioorg Med. Chem. Lett.</source> <volume>26</volume>, <fpage>1224</fpage>&#x2013;<lpage>1228</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmcl.2016.01.032</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gaulton</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bellis</surname>
<given-names>L. J.</given-names>
</name>
<name>
<surname>Bento</surname>
<given-names>A. P.</given-names>
</name>
<name>
<surname>Chambers</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Davies</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Hersey</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2012</year>). <article-title>ChEMBL: A large-scale bioactivity database for drug discovery</article-title>. <source>Nucleic Acids Res.</source> <volume>40</volume>, <fpage>D1100</fpage>&#x2013;<lpage>D1107</lpage>. <pub-id pub-id-type="doi">10.1093/nar/gkr777</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hanley</surname>
<given-names>J. A.</given-names>
</name>
<name>
<surname>McNeil</surname>
<given-names>B. J.</given-names>
</name>
</person-group> (<year>1983</year>). <article-title>A method of comparing the areas under receiver operating characteristic curves derived from the same cases</article-title>. <source>Radiology</source> <volume>148</volume>, <fpage>839</fpage>&#x2013;<lpage>843</lpage>. <pub-id pub-id-type="doi">10.1148/radiology.148.3.6878708</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Chan</surname>
<given-names>G. K.</given-names>
</name>
<name>
<surname>Chang</surname>
<given-names>J. H.</given-names>
</name>
<name>
<surname>Do</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Drummond</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2015</year>). <article-title>Discovery of 3, 5-substituted 6-azaindazoles as potent pan-Pim inhibitors</article-title>. <source>Bioorg Med. Chem. Lett.</source> <volume>25</volume>, <fpage>5258</fpage>&#x2013;<lpage>5264</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmcl.2015.09.052</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yuan</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Xu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>W.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Identification of pim-1 kinase inhibitors by pharmacophore model, molecular docking-based virtual screening, and biological evaluation</article-title>. <source>Curr. Comput. Aided. Drug. Des.</source> <volume>18</volume>, <fpage>240</fpage>&#x2013;<lpage>246</lpage>. <pub-id pub-id-type="doi">10.2174/1573409918666220427120524</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ibrahim</surname>
<given-names>M. H.</given-names>
</name>
<name>
<surname>Harras</surname>
<given-names>M. F.</given-names>
</name>
<name>
<surname>Mostafa</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Mohyeldin</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Kamaly</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Altwaijry</surname>
<given-names>N.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Development of novel cyanopyridines as PIM-1 kinase inhibitors with potent anti-prostate cancer activity: Synthesis, biological evaluation, nanoparticles formulation and molecular dynamics simulation</article-title>. <source>Bioorg Chem.</source> <volume>129</volume>, <fpage>106122</fpage>. <pub-id pub-id-type="doi">10.1016/j.bioorg.2022.106122</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kursa</surname>
<given-names>M. B.</given-names>
</name>
<name>
<surname>Jankowski</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Rudnicki</surname>
<given-names>W. R.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Boruta&#x2014;a system for feature selection</article-title>. <source>Fundam. Inf.</source> <volume>101</volume>, <fpage>271</fpage>&#x2013;<lpage>285</lpage>. <pub-id pub-id-type="doi">10.3233/fi-2010-288</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Le</surname>
<given-names>B. T.</given-names>
</name>
<name>
<surname>Kumarasiri</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Adams</surname>
<given-names>J. R.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Milne</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Sykes</surname>
<given-names>M. J.</given-names>
</name>
<etal/>
</person-group> (<year>2015</year>). <article-title>Targeting pim kinases for cancer treatment: Opportunities and challenges</article-title>. <source>Future. Med. Chem.</source> <volume>7</volume>, <fpage>35</fpage>&#x2013;<lpage>53</lpage>. <pub-id pub-id-type="doi">10.4155/fmc.14.145</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Loveland</surname>
<given-names>B. E.</given-names>
</name>
<name>
<surname>Xing</surname>
<given-names>P. X.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>Anti-Pim-1 mAb inhibits activation and proliferation of T lymphocytes and prolongs mouse skin allograft survival</article-title>. <source>Cell Immunol.</source> <volume>272</volume>, <fpage>87</fpage>&#x2013;<lpage>93</lpage>. <pub-id pub-id-type="doi">10.1016/j.cellimm.2011.09.002</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Fan</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>T.</given-names>
</name>
<etal/>
</person-group> (<year>2016</year>). <article-title>Synthesis and biological evaluation of quinoline derivatives as potential anti-prostate cancer agents and Pim-1 kinase inhibitors</article-title>. <source>Bioorg Med. Chem.</source> <volume>24</volume> (<issue>8</issue>), <fpage>1889</fpage>&#x2013;<lpage>1897</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmc.2016.03.016</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>Y. Y.</given-names>
</name>
<name>
<surname>Popivanova</surname>
<given-names>B. K.</given-names>
</name>
<name>
<surname>Nagai</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Ishikura</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Fujii</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Mukaida</surname>
<given-names>N.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>Pim-3 a proto-oncogene with serine/threonine kinase activity, is aberrantly expressed in human pancreatic cancer and phosphorylates bad to block bad-mediated apoptosis in human pancreatic cancer cell lines</article-title>. <source>Cancer Res.</source> <volume>66</volume>, <fpage>6741</fpage>&#x2013;<lpage>6747</lpage>. <pub-id pub-id-type="doi">10.1158/0008-5472.can-05-4272</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lipi&#x144;ski</surname>
<given-names>P. F. J.</given-names>
</name>
<name>
<surname>Szurmak</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>SCRAMBLE&#x2019;N&#x2019;GAMBLE: A tool for fast and facile generation of random data for statistical evaluation of QSAR models</article-title>. <source>Chem. Pap.</source> <volume>71</volume>, <fpage>2217</fpage>&#x2013;<lpage>2232</lpage>. <pub-id pub-id-type="doi">10.1007/s11696-017-0215-7</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Minerali</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Foil</surname>
<given-names>D. H.</given-names>
</name>
<name>
<surname>Zorn</surname>
<given-names>K. M.</given-names>
</name>
<name>
<surname>Lane</surname>
<given-names>T. R.</given-names>
</name>
<name>
<surname>Ekins</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Comparing machine learning algorithms for predicting drug-induced liver injury (DILI)</article-title>. <source>Mol. Pharm.</source> <volume>17</volume>, <fpage>2628</fpage>&#x2013;<lpage>2637</lpage>. <pub-id pub-id-type="doi">10.1021/acs.molpharmaceut.0c00326</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Mitchell</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>1997</year>). <source>Machine learning</source>. <publisher-loc>New York</publisher-loc>: <publisher-name>McGraw-Hill</publisher-name>.</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Narlik-Grassow</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Blanco-Aparicio</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Carnero</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>The PIM family of serine/threonine kinases in cancer</article-title>. <source>Med. Res. Rev.</source> <volume>34</volume>, <fpage>136</fpage>&#x2013;<lpage>159</lpage>.</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nawijn</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Alendar</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Berns</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>For better or for worse: The role of pim oncogenes in tumorigenesis</article-title>. <source>Nat. Rev. Cancer</source> <volume>11</volume> (<issue>11</issue>), <fpage>23</fpage>&#x2013;<lpage>34</lpage>. <pub-id pub-id-type="doi">10.1038/nrc2986</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nonga</surname>
<given-names>O. E.</given-names>
</name>
<name>
<surname>Lavogina</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Enkvist</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Kestav</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Chaikuad</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Dixon-Clarke</surname>
<given-names>S. E.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Crystal structure-guided design of bisubstrate inhibitors and photoluminescent probes for protein kinases of the PIM family</article-title>. <source>Molecules</source> <volume>26</volume>, <fpage>4353</fpage>. <pub-id pub-id-type="doi">10.3390/molecules26144353</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ogawa</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Yuki</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Tanaka</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Insights from Pim1 structure for anti-cancer drug design</article-title>. <source>Expert Opin. Drug Discov.</source> <volume>7</volume>, <fpage>1177</fpage>&#x2013;<lpage>1192</lpage>. <pub-id pub-id-type="doi">10.1517/17460441.2012.727394</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Park</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Jeon</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Choi</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Hong</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Structure-based virtual screening and de novo design of PIM1 inhibitors with anticancer activity from natural products</article-title>. <source>Pharm. (Basel)</source> <volume>14</volume>, <fpage>275</fpage>. <pub-id pub-id-type="doi">10.3390/ph14030275</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Park</surname>
<given-names>S. H.</given-names>
</name>
<name>
<surname>Goo</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Jo</surname>
<given-names>C. H.</given-names>
</name>
</person-group> (<year>2004</year>). <article-title>Receiver operating characteristic (ROC) curve: Practical review for radiologists</article-title>. <source>Korean J. Radiol.</source> <volume>5</volume>, <fpage>11</fpage>&#x2013;<lpage>18</lpage>. <pub-id pub-id-type="doi">10.3348/kjr.2004.5.1.11</pub-id>
</citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Parrinello</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Rahman</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>1981</year>). <article-title>Polymorphic transitions in single crystals: A new molecular dynamics method</article-title>. <source>J. Appl. Phys.</source> <volume>52</volume>, <fpage>7182</fpage>&#x2013;<lpage>7190</lpage>. <pub-id pub-id-type="doi">10.1063/1.328693</pub-id>
</citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pogacic</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Bullock</surname>
<given-names>A. N.</given-names>
</name>
<name>
<surname>Fedorov</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Filippakopoulos</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Gasser</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Biondi</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2007</year>). <article-title>Structural analysis identifies imidazo[1, 2-b]pyridazines as PIM kinase inhibitors with <italic>in vitro</italic> antileukemic activity</article-title>. <source>Cancer Res.</source> <volume>67</volume>, <fpage>6916</fpage>&#x2013;<lpage>6924</lpage>. <pub-id pub-id-type="doi">10.1158/0008-5472.can-07-0320</pub-id>
</citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ponzoni</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Sebasti&#xe1;n-P&#xe9;rez</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Mart&#xed;nez</surname>
<given-names>M. J.</given-names>
</name>
<name>
<surname>Roca</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Cruz P&#xe9;rez</surname>
<given-names>C. D.</given-names>
</name>
<name>
<surname>Cravero</surname>
<given-names>F.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>QSAR classification models for predicting the activity of inhibitors of beta-secretase (BACE1) associated with alzheimer&#x27;s disease</article-title>. <source>Sci. Rep.</source> <volume>9</volume>, <fpage>9102</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-019-45522-3</pub-id>
</citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pronk</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>P&#xe1;ll</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Schulz</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Larsson</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Bjelkmar</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Apostolov</surname>
<given-names>R.</given-names>
</name>
<etal/>
</person-group> (<year>2013</year>). <article-title>Gromacs 4.5: A high-throughput and highly parallel open source molecular simulation toolkit</article-title>. <source>Bioinformatics</source> <volume>29</volume>, <fpage>845</fpage>&#x2013;<lpage>854</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btt055</pub-id>
</citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Qian</surname>
<given-names>K. C.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Hickey</surname>
<given-names>E. R.</given-names>
</name>
<name>
<surname>Studts</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Barringer</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Peng</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2005</year>). <article-title>Structural basis of constitutive activity and a unique nucleotide binding mode of human Pim-1 kinase</article-title>. <source>J. Biol. Chem.</source> <volume>280</volume>, <fpage>6130</fpage>&#x2013;<lpage>6137</lpage>. <pub-id pub-id-type="doi">10.1074/jbc.m409123200</pub-id>
</citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Rakhimbekova</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Madzhidov</surname>
<given-names>T. I.</given-names>
</name>
<name>
<surname>Nugmanov</surname>
<given-names>R. I.</given-names>
</name>
<name>
<surname>Gimadiev</surname>
<given-names>T. R.</given-names>
</name>
<name>
<surname>Baskin</surname>
<given-names>I. I.</given-names>
</name>
<name>
<surname>Varnek</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Comprehensive analysis of applicability domains of QSPR models for chemical reactions</article-title>. <source>Int. J. Mol. Sci.</source> <volume>21</volume>, <fpage>5542</fpage>. <pub-id pub-id-type="doi">10.3390/ijms21155542</pub-id>
</citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>R&#xfc;cker</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>R&#xfc;cker</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Meringer</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>y-Randomization and its variants in QSPR/QSAR</article-title>. <source>J. Chem. Inf. Model</source> <volume>47</volume>, <fpage>2345</fpage>&#x2013;<lpage>2357</lpage>. <pub-id pub-id-type="doi">10.1021/ci700157b</pub-id>
</citation>
</ref>
<ref id="B42">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shaik</surname>
<given-names>N. A.</given-names>
</name>
<name>
<surname>Nasser</surname>
<given-names>K. K.</given-names>
</name>
<name>
<surname>Alruwaili</surname>
<given-names>M. M.</given-names>
</name>
<name>
<surname>Alallasi</surname>
<given-names>S. R.</given-names>
</name>
<name>
<surname>Elango</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Banaganapalli</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Molecular modelling and dynamic simulations of sequestosome 1 (SQSTM1) missense mutations linked to Paget disease of bone</article-title>. <source>J. Biomol. Struct. Dyn.</source> <volume>39</volume>, <fpage>2873</fpage>&#x2013;<lpage>2884</lpage>. <pub-id pub-id-type="doi">10.1080/07391102.2020.1758212</pub-id>
</citation>
</ref>
<ref id="B43">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Shaik</surname>
<given-names>N. A.</given-names>
</name>
<name>
<surname>Saud Al-Saud</surname>
<given-names>N. B.</given-names>
</name>
<name>
<surname>Abdulhamid Aljuhani</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Jamil</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Alnuman</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Aljeaid</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>Structural characterization and conformational dynamics of alpha-1 antitrypsin pathogenic variants causing alpha-1-antitrypsin deficiency</article-title>. <source>Front. Mol. Biosci.</source> <volume>9</volume>, <fpage>1051511</fpage>. <pub-id pub-id-type="doi">10.3389/fmolb.2022.1051511</pub-id>
</citation>
</ref>
<ref id="B44">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sushko</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Novotarskyi</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>K&#xf6;rner</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Pandey</surname>
<given-names>A. K.</given-names>
</name>
<name>
<surname>Cherkasov</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2010</year>). <article-title>Applicability domains for classification problems: Benchmarking of distance to models for ames mutagenicity set</article-title>. <source>J. Chem. Inf. Model.</source> <volume>50</volume>, <fpage>2094</fpage>&#x2013;<lpage>2111</lpage>. <pub-id pub-id-type="doi">10.1021/ci100253r</pub-id>
</citation>
</ref>
<ref id="B45">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Trott</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>OlsonAutoDock</surname>
<given-names>A. J.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>AutoDock Vina: Improving the speed and accuracy of docking with a new scoring function, efficient optimization, and multithreading</article-title>. <source>J. Comput. Chem.</source> <volume>31</volume>, <fpage>455</fpage>&#x2013;<lpage>461</lpage>. <pub-id pub-id-type="doi">10.1002/jcc.21334</pub-id>
</citation>
</ref>
<ref id="B46">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tsuganezawa</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Watanabe</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Parker</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Yuki</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Taruya</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Nakagawa</surname>
<given-names>Y.</given-names>
</name>
<etal/>
</person-group> (<year>2012</year>). <article-title>A novel Pim-1 kinase inhibitor targeting residues that bind the substrate peptide</article-title>. <source>J. Mol. Biol.</source> <volume>417</volume>, <fpage>240</fpage>&#x2013;<lpage>252</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmb.2012.01.036</pub-id>
</citation>
</ref>
<ref id="B47">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tursynbay</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Tokay</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Zhumadilov</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2016</year>). <article-title>Pim-1 kinase as cancer drug target: An update</article-title>. <source>Biomed. Rep.</source> <volume>4</volume>, <fpage>140</fpage>&#x2013;<lpage>146</lpage>. <pub-id pub-id-type="doi">10.3892/br.2015.561</pub-id>
</citation>
</ref>
<ref id="B48">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Vivek</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bharti</surname>
<given-names>S. K.</given-names>
</name>
<name>
<surname>Kumar Budhwani</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>3D-QSAR and virtual screening studies of thiazolidine-2, 4-dione analogs: Validation of experimental inhibitory potencies towards PIM-1 kinase</article-title>. <source>J. Mol. Struct.</source> <volume>1133</volume>, <fpage>278</fpage>&#x2013;<lpage>293</lpage>. <pub-id pub-id-type="doi">10.1016/j.molstruc.2016.12.006</pub-id>
</citation>
</ref>
<ref id="B49">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Voulgaris</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Magoulas</surname>
<given-names>G. D.</given-names>
</name>
</person-group> (<year>2008</year>). &#x201c;<article-title>Extensions of the k nearest neighbour methods for classification problems</article-title>,&#x201d; in <conf-name>Proceedings of the 26th IASTED International Conference on Artificial Intelligence and Applications (AIA &#x2018;08)</conf-name>, <conf-loc>Anaheim, CA</conf-loc>, <conf-date>February, 2008</conf-date> (<publisher-name>ACTA Press</publisher-name>), <fpage>23</fpage>&#x2013;<lpage>28</lpage>.</citation>
</ref>
<ref id="B50">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Magnuson</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Pastor</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Fan</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Tsui</surname>
<given-names>V.</given-names>
</name>
<etal/>
</person-group> (<year>2013</year>). <article-title>Discovery of novel pyrazolo[1, 5-a]pyrimidines as potent pan-Pim inhibitors by structure- and property-based drug design</article-title>. <source>Bioorg Med. Chem. Lett.</source> <volume>23</volume>, <fpage>3149</fpage>&#x2013;<lpage>3153</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmcl.2013.04.020</pub-id>
</citation>
</ref>
<ref id="B51">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Warfel</surname>
<given-names>N. A.</given-names>
</name>
<name>
<surname>Kraft</surname>
<given-names>A. S.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>PIM kinase (and Akt) biology and signaling in tumors</article-title>. <source>Pharmacol. Ther.</source> <volume>151</volume>, <fpage>41</fpage>&#x2013;<lpage>49</lpage>. <pub-id pub-id-type="doi">10.1016/j.pharmthera.2015.03.001</pub-id>
</citation>
</ref>
<ref id="B52">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wurz</surname>
<given-names>R. P.</given-names>
</name>
<name>
<surname>Pettus</surname>
<given-names>L. H.</given-names>
</name>
<name>
<surname>Jackson</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>H. L.</given-names>
</name>
<name>
<surname>Herberich</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group> (<year>2015</year>). <article-title>The discovery and optimization of aminooxadiazoles as potent Pim kinase inhibitors</article-title>. <source>Bioorg Med. Chem. Lett.</source> <volume>25</volume> (<issue>4</issue>), <fpage>847</fpage>&#x2013;<lpage>855</lpage>. <pub-id pub-id-type="doi">10.1016/j.bmcl.2014.12.067</pub-id>
</citation>
</ref>
<ref id="B53">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xia</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Knaak</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Beharry</surname>
<given-names>Z. M.</given-names>
</name>
<name>
<surname>McInnes</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>W.</given-names>
</name>
<etal/>
</person-group> (<year>2009</year>). <article-title>Synthesis and evaluation of novel inhibitors of Pim-1 and Pim-2 protein kinases</article-title>. <source>J. Med. Chem.</source> <volume>52</volume>, <fpage>74</fpage>&#x2013;<lpage>86</lpage>. <pub-id pub-id-type="doi">10.1021/jm800937p</pub-id>
</citation>
</ref>
<ref id="B54">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yap</surname>
<given-names>C. W.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>PaDEL-descriptor: An open source software to calculate molecular descriptors and f ingerprints</article-title>. <source>J. Comput. Chem.</source> <volume>32</volume>, <fpage>1466</fpage>&#x2013;<lpage>1474</lpage>. <pub-id pub-id-type="doi">10.1002/jcc.21707</pub-id>
</citation>
</ref>
<ref id="B55">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhao</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Qiu</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>PIM1: A promising target in patients with triple-negative breast cancer</article-title>. <source>Med. Oncol.</source> <volume>34</volume>, <fpage>142</fpage>. <pub-id pub-id-type="doi">10.1007/s12032-017-0998-y</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>