<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Psychiatry</journal-id>
<journal-title>Frontiers in Psychiatry</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Psychiatry</abbrev-journal-title>
<issn pub-type="epub">1664-0640</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpsyt.2023.1266548</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Psychiatry</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Prediction of patient admission and readmission in adults from a Colombian cohort with bipolar disorder using artificial intelligence</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes"><name><surname>Palacios-Ariza</surname> <given-names>Mar&#x00ED;a Alejandra</given-names></name><xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
</contrib>
<contrib contrib-type="author" equal-contrib="yes"><name><surname>Morales-Mendoza</surname> <given-names>Esteban</given-names></name><xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="author-notes" rid="fn0001"><sup>&#x2020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes"><name><surname>Murcia</surname> <given-names>Jossie</given-names></name><xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="author-notes" rid="fn0001"><sup>&#x2020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes"><name><surname>Arias-Duarte</surname> <given-names>Rafael</given-names></name><xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<xref ref-type="author-notes" rid="fn0001"><sup>&#x2020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes"><name><surname>Lara-Castellanos</surname> <given-names>Germ&#x00E1;n</given-names></name><xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<xref ref-type="author-notes" rid="fn0001"><sup>&#x2020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes"><name><surname>Cely-Jim&#x00E9;nez</surname> <given-names>Andr&#x00E9;s</given-names></name><xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<xref ref-type="author-notes" rid="fn0001"><sup>&#x2020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes"><name><surname>Rinc&#x00F3;n-Acu&#x00F1;a</surname> <given-names>Juan Carlos</given-names></name><xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
<xref ref-type="author-notes" rid="fn0002"><sup>&#x2021;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2592301/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes"><name><surname>Ara&#x00FA;zo-Bravo</surname> <given-names>Marcos J.</given-names></name><xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<xref ref-type="aff" rid="aff6"><sup>6</sup></xref>
<xref ref-type="aff" rid="aff7"><sup>7</sup></xref>
<xref ref-type="aff" rid="aff8"><sup>8</sup></xref>
<xref ref-type="author-notes" rid="fn0002"><sup>&#x2021;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/142192/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes"><name><surname>McDouall</surname> <given-names>Jorge</given-names></name><xref ref-type="aff" rid="aff9"><sup>9</sup></xref>
<xref ref-type="author-notes" rid="fn0002"><sup>&#x2021;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Unidad de Investigaci&#x00F3;n, Fundaci&#x00F3;n Universitaria Sanitas, Psicopatolog&#x00ED;a y Sociedad Research Group</institution>, <addr-line>Bogot&#x00E1;</addr-line>, <country>Colombia</country></aff>
<aff id="aff2"><sup>2</sup><institution>Fundaci&#x00F3;n Universitaria Sanitas, Gerencia y Gesti&#x00F3;n Sanitaria Research Group, Instituto de Gerencia y Gesti&#x00F3;n Sanitaria (IGGS)</institution>, <addr-line>Bogot&#x00E1;</addr-line>, <country>Colombia</country></aff>
<aff id="aff3"><sup>3</sup><institution>Psicopatolog&#x00ED;a y Sociedad Research Group, Facultad de Medicina, Fundaci&#x00F3;n Universitaria Sanitas</institution>, <addr-line>Bogot&#x00E1;</addr-line>, <country>Colombia</country></aff>
<aff id="aff4"><sup>4</sup><institution>Keralty</institution>, <addr-line>Bogot&#x00E1;</addr-line>, <country>Colombia</country></aff>
<aff id="aff5"><sup>5</sup><institution>University of Santander - UDES</institution>, <addr-line>Bucaramanga</addr-line>, <country>Colombia</country></aff>
<aff id="aff6"><sup>6</sup><institution>Computational Biology and Systems Biomedicine, Biodonostia Health Research Institute</institution>, <addr-line>San Sebasti&#x00E1;n</addr-line>, <country>Spain</country></aff>
<aff id="aff7"><sup>7</sup><institution>Ikerbasque, Basque Foundation for Science</institution>, <addr-line>Bilbao</addr-line>, <country>Spain</country></aff>
<aff id="aff8"><sup>8</sup><institution>Department of Cell Biology and Histology, Faculty of Medicine and Nursing, University of Basque Country (UPV/EHU)</institution>, <addr-line>Leioa</addr-line>, <country>Spain</country></aff>
<aff id="aff9"><sup>9</sup><institution>Sanitas Crea Research Group, Fundaci&#x00F3;n Universitaria Sanitas</institution>, <addr-line>Bogot&#x00E1;</addr-line>, <country>Colombia</country></aff>
<author-notes>
<fn fn-type="edited-by" id="fn0003">
<p>Edited by: Steven Fernandes, Creighton University, United States</p>
</fn>
<fn fn-type="edited-by" id="fn0004">
<p>Reviewed by: Xinyuan Yan, University of Minnesota Twin Cities, United States; Ebrahim Ghaderpour, Sapienza University of Rome, Italy</p>
</fn>
<corresp id="c001">&#x002A;Correspondence: Mar&#x00ED;a Alejandra Palacios-Ariza, <email>mapalaciosar@unisanitas.edu.co</email></corresp>
<fn fn-type="equal" id="fn0001">
<p><sup>&#x2020;</sup>These authors have contributed equally to this work</p>
</fn>
<fn fn-type="equal" id="fn0002">
<p><sup>&#x2021;</sup>These authors share senior authorship</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>21</day>
<month>12</month>
<year>2023</year>
</pub-date>
<pub-date pub-type="collection">
<year>2023</year>
</pub-date>
<volume>14</volume>
<elocation-id>1266548</elocation-id>
<history>
<date date-type="received">
<day>25</day>
<month>07</month>
<year>2023</year>
</date>
<date date-type="accepted">
<day>30</day>
<month>11</month>
<year>2023</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2023 Palacios-Ariza, Morales-Mendoza, Murcia, Arias-Duarte, Lara-Castellanos, Cely-Jim&#x00E9;nez, Rinc&#x00F3;n-Acu&#x00F1;a, Ara&#x00FA;zo-Bravo and McDouall.</copyright-statement>
<copyright-year>2023</copyright-year>
<copyright-holder>Palacios-Ariza, Morales-Mendoza, Murcia, Arias-Duarte, Lara-Castellanos, Cely-Jim&#x00E9;nez, Rinc&#x00F3;n-Acu&#x00F1;a, Ara&#x00FA;zo-Bravo and McDouall</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<sec id="sec1">
<title>Introduction</title>
<p>Bipolar disorder (BD) is a chronically progressive mental condition, associated with a reduced quality of life and greater disability. Patient admissions are preventable events with a considerable impact on global functioning and social adjustment. While machine learning (ML) approaches have proven prediction ability in other diseases, little is known about their utility to predict patient admissions in this pathology.</p>
</sec>
<sec id="sec2">
<title>Aim</title>
<p>To develop prediction models for hospital admission/readmission within 5&#x2009;years of diagnosis in patients with BD using ML techniques.</p>
</sec>
<sec id="sec3">
<title>Methods</title>
<p>The study utilized data from patients diagnosed with BD in a major healthcare organization in Colombia. Candidate predictors were selected from Electronic Health Records (EHRs) and included sociodemographic and clinical variables. ML algorithms, including Decision Trees, Random Forests, Logistic Regressions, and Support Vector Machines, were used to predict patient admission or readmission. Survival models, including a penalized Cox Model and Random Survival Forest, were used to predict time to admission and first readmission. Model performance was evaluated using accuracy, precision, recall, F1 score, area under the receiver operating characteristic curve (AUC) and concordance index.</p>
</sec>
<sec id="sec4">
<title>Results</title>
<p>The admission dataset included 2,726 BD patients, with 354 admissions, while the readmission dataset included 352 patients, with almost half being readmitted. The best-performing model for predicting admission was the Random Forest, with an accuracy score of 0.951 and an AUC of 0.98. The variables with the greatest predictive power in the Recursive Feature Elimination (RFE) importance analysis were the number of psychiatric emergency visits, the number of outpatient follow-up appointments and age. Survival models showed similar results, with the Random Survival Forest performing best, achieving an AUC of 0.95. However, the prediction models for patient readmission had poorer performance, with the Random Forest model being again the best performer but with an AUC below 0.70.</p>
</sec>
<sec id="sec5">
<title>Conclusion</title>
<p>ML models, particularly the Random Forest model, outperformed traditional statistical techniques for admission prediction. However, readmission prediction models had poorer performance. This study demonstrates the potential of ML techniques in improving prediction accuracy for BD patient admissions.</p>
</sec>
</abstract>
<kwd-group>
<kwd>bipolar disorder</kwd>
<kwd>electronic health records</kwd>
<kwd>machine learning</kwd>
<kwd>patient admission</kwd>
<kwd>patient readmission</kwd>
<kwd>risk factors</kwd>
</kwd-group>
<counts>
<fig-count count="4"/>
<table-count count="2"/>
<equation-count count="0"/>
<ref-count count="44"/>
<page-count count="13"/>
<word-count count="7002"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Computational Psychiatry</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec6"><label>1</label>
<title>Introduction</title>
<p>Bipolar disorder (BD) is a chronically progressive mental disorder with a prevalence that ranges from 1.1 to 2.4% (<xref ref-type="bibr" rid="ref1">1</xref>, <xref ref-type="bibr" rid="ref2">2</xref>). BD is classified as type I if the patient has presented at least one manic episode, with or without depressive episodes, and as type II in the presence of at least one hypomanic episode, with no full manic episodes, and one major depressive episode (<xref ref-type="bibr" rid="ref3">3</xref>). This condition is associated with a reduced quality of life and greater disability. Patients have been shown to have lower incomes, higher financial burdens, issues with social interactions and a greater overall frequency of use of health services (<xref ref-type="bibr" rid="ref4 ref5 ref6 ref7">4&#x2013;7</xref>). Furthermore, significant suicide attempt rates have been reported in patients with BD type I (36.3%) and BD type II (32.4%) (<xref ref-type="bibr" rid="ref8">8</xref>, <xref ref-type="bibr" rid="ref9">9</xref>). Additionally, manic episodes give rise to destructive and reckless behavior secondary to an unstable mood, which in the long term can degrade cognitive functioning, interpersonal relationships, global functioning and social adjustment (<xref ref-type="bibr" rid="ref10">10</xref>, <xref ref-type="bibr" rid="ref11">11</xref>).</p>
<p>Patient admissions are preventable events with a considerable impact on healthcare costs and a key quality metric for health systems around the world (<xref ref-type="bibr" rid="ref12 ref13 ref14">12&#x2013;14</xref>). Among psychiatric patients, whose disorders are characterized by chronicity and high recurrence rates, readmissions are of particular concern. The period immediately after a hospitalization is known to be a period of high risk for outcomes such as suicide and substance abuse relapse. Patient admission and readmission due to BD places a significant financial burden on medical services and caregivers (<xref ref-type="bibr" rid="ref2">2</xref>, <xref ref-type="bibr" rid="ref15">15</xref>). Studies have reported relapse rates of up to 50% after 2&#x2009;years in individuals receiving adequate psychopharmacological care (<xref ref-type="bibr" rid="ref16">16</xref>). Self-monitoring data from the Sanitas Healthcare Management Organization (HMO) in Colombia show that in 2017, 319 patients with a BD diagnosis had a total of 427 hospital admissions (27% of the overall number of readmissions) with a mean hospital stay duration of 19&#x2009;days.</p>
<p>The prediction of risk for patient admission and readmission may aid in disease management as well as in reducing the economic and social burdens caused by BD. Despite the considerable number of publications in the field of patient admission prediction, most studies address this issue in patients with non-psychiatric illness (<xref ref-type="bibr" rid="ref17 ref18 ref19 ref20 ref21">17&#x2013;21</xref>), substance abuse disorders, schizophrenia or postpartum depression (<xref ref-type="bibr" rid="ref22 ref23 ref24">22&#x2013;24</xref>). Prior work by Rotenberg et al. aimed to predict depressive relapses in patients with bipolar disorder using machine learning techniques, achieving F measures as high as 0.993 for a random forest model (<xref ref-type="bibr" rid="ref25">25</xref>). Although certainly related, depressive relapses are only one potential form of relapse of BD and may indeed be less likely to require hospitalization than manic episodes. This is why this study developed prediction models for hospital admission/readmission within 5&#x2009;years of diagnosis in patients with BD using ML techniques. With clinical data that is readily available in electronic health records, we show that random forest models are good predictors of both admission and time to admission in patients with BD.</p>
</sec>
<sec sec-type="materials|methods" id="sec7"><label>2</label>
<title>Materials and methods</title>
<sec id="sec8"><label>2.1</label>
<title>Study participants and design</title>
<p>Sanitas is a major Healthcare Management Organization (HMO) in Colombia with over 5 million patients under its care. Data used in this study were obtained from patients who received an incident diagnosis of BD (including both type I and type II patients) during outpatient visits to any of the Sanitas EPS healthcare facilities in Colombia during 2016. The patients were followed for 5&#x2009;years after diagnosis until December 31st, 2020. The Sanitas EPS network includes high-complexity centers which offer inpatient mental health services, day hospitals, and priority and general psychiatric outpatient appointments. Patients were retrospectively included if they had an International Classification of Diseases 10 diagnostic code of F31 assigned to them at any of their outpatient visits during the study period. Subjects with a BD diagnosis were excluded if they required long-term hospitalization due to functioning or psychotic symptoms, or if their social context (poor social support) required them to remain hospitalized in mental health units despite their psychiatric condition no longer requiring hospitalization. This study followed the guidelines of the Colombian Ministry of Health resolution 8,430 of 1993, as well as the World Medical Association&#x2019;s Declaration of Helsinki in its 2013 version, and the Council for International Organizations of Medical Sciences&#x2019; International Ethical Guidelines for Health-related Research Involving Humans. The protocol for this study was approved by the Research Ethics Committee of the Fundaci&#x00F3;n Universitaria Sanitas (CEIFUS 341&#x2013;19).</p>
<p>Starting from an overall HMO patient cohort of over 3 million patients over the age of 18, we selected patients with an incident diagnosis of BD (diagnostic code F31), which totaled 2,726 patients. Of these, 354 patients had at least one psychiatric ward inpatient admission registered over 5&#x2009;years. A complete record for the admission was available for 352 patients. These 352 patients comprised the admission dataset. Readmission occurred in 165 patients in whom a complete record for the readmission was available. The readmission dataset was obtained from these 165 patients (<xref ref-type="fig" rid="fig1">Figure 1</xref>).</p>
<fig position="float" id="fig1"><label>Figure 1</label>
<caption>
<p>Patient flowchart.</p>
</caption>
<graphic xlink:href="fpsyt-14-1266548-g001.tif"/>
</fig>
</sec>
<sec id="sec9"><label>2.2</label>
<title>Outcome definitions</title>
<p>The primary outcome in this study was a composite of hospital admission or readmission during the study period, although these outcomes were considered separately in model construction and selection. We selected this composite outcome as it is a key indicator of disease status in BD. An initial group of models aimed to predict a binary outcome of admission/readmission during the observation period, and a second group of survival models aimed to predict time to admission/readmission.</p>
</sec>
<sec id="sec10"><label>2.3</label>
<title>Candidate predictors</title>
<p>Candidate predictors for ML models were selected from features available in the standard electronic health record (EHR) and making use of established knowledge regarding risk factors for the study outcome. The initial set of extracted predictors included sociodemographic variables, and variables related to the psychiatric and medical history. Specifically, for the outcome of patient readmission a group of variables describing the characteristics of the prior hospitalization was included.</p>
<p>A total of 47 candidate variables were considered initially for the outcome of admission. However, since not all variables were plausibly related to the outcome of either admission or readmission (e.g., type of discharge cannot be used as a predictor), this led to variable exclusion and a reduced set of 29 variables. After completeness analysis, we excluded variables with missing data in more than 35% of the subjects, further reducing the set to 18 variables. The readmission dataset was refined from the same starting pool of variables. After an initial analysis, we were left with a reduced set of 33 variables (note that variables like &#x201C;type of discharge&#x201D; for the first admission can now be included). After completeness analysis, we were left with 21 predictor variables (<xref ref-type="supplementary-material" rid="SM1">Supplementary Table S1</xref>).</p>
<p>Variables with multiple categories were dichotomized into the presence or absence of their most frequent value prior to inclusion in the models but are presented in <xref ref-type="table" rid="tab1">Table 1</xref> as extracted from the EHRs. Continuous data were normalized/standardized and analyzed using correlation matrices. Variables were normalized by subtracting the minimum value of the variable and dividing the result by its range (difference between maximum and minimum). Standardized variables were the result of subtracting by the variables mean and dividing the result by their standard deviation. Although we planned to merge variables with correlations exceeding a pre-established threshold of 0.8, no such correlations were identified. Feature engineering made use of two techniques: Recursive Feature Elimination (RFE) and Sequential Feature Selection (SFS) (<xref ref-type="fig" rid="fig2">Figure 2</xref>). The RFE used a random forest to select features based on the top F1 scores, and the SFS used a sequential method to obtain a reduced set which minimized standard error and dimensionality (<xref ref-type="bibr" rid="ref26">26</xref>, <xref ref-type="bibr" rid="ref27">27</xref>).</p>
<table-wrap position="float" id="tab1"><label>Table 1</label>
<caption>
<p>Patient characteristics according to patient admission/readmission stratification.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top" rowspan="2">Patient admission set of variables (18 variables)</th>
<th align="center" valign="top">Total (<italic>N</italic>&#x2009;=&#x2009;2,726)</th>
<th align="center" valign="top">Admission (<italic>n</italic>&#x2009;=&#x2009;354)</th>
<th align="center" valign="top">No admission (<italic>n</italic>&#x2009;=&#x2009;2,372)</th>
<th align="center" valign="top" rowspan="2"><italic>p-</italic>value</th>
</tr>
<tr>
<th align="center" valign="top" colspan="3">n (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="bottom">Female</td>
<td align="center" valign="bottom">1747 (64.09)</td>
<td align="center" valign="bottom">229 (13.11)</td>
<td align="center" valign="bottom">1,518 (86.89)</td>
<td align="center" valign="bottom">0.846</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Marital status</td>
</tr>
<tr>
<td align="left" valign="bottom">Married</td>
<td align="center" valign="bottom">571 (20.95)</td>
<td align="center" valign="bottom">71 (12.43)</td>
<td align="center" valign="bottom">500 (87.57)</td>
<td align="center" valign="middle" rowspan="5">0.253&#x002A;</td>
</tr>
<tr>
<td align="left" valign="bottom">Divorced</td>
<td align="center" valign="bottom">27 (0.99)</td>
<td align="center" valign="bottom">6 (22.22)</td>
<td align="center" valign="bottom">21 (77.78)</td>
</tr>
<tr>
<td align="left" valign="bottom">Single</td>
<td align="center" valign="bottom">1902 (69.77)</td>
<td align="center" valign="bottom">256 (13.46)</td>
<td align="center" valign="bottom">1,646 (86.54)</td>
</tr>
<tr>
<td align="left" valign="bottom">Common law marriage</td>
<td align="center" valign="bottom">209 (7.67)</td>
<td align="center" valign="bottom">20 (9.57)</td>
<td align="center" valign="bottom">189 (90.43)</td>
</tr>
<tr>
<td align="left" valign="bottom">Widowed</td>
<td align="center" valign="bottom">17 (0.62)</td>
<td align="center" valign="bottom">1 (5.88)</td>
<td align="center" valign="bottom">16 (94.12)</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Healthcare regime</td>
</tr>
<tr>
<td align="left" valign="bottom">Contributive</td>
<td align="center" valign="bottom">2,439 (89.47)</td>
<td align="center" valign="bottom">303 (12.42)</td>
<td align="center" valign="bottom">2,136 (87.58)</td>
<td align="center" valign="middle" rowspan="3">
<bold>&#x003C;0.001</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">Premium plan</td>
<td align="center" valign="bottom">189 (6.93)</td>
<td align="center" valign="bottom">42 (22.22)</td>
<td align="center" valign="bottom">147 (77.77)</td>
</tr>
<tr>
<td align="left" valign="bottom">Subsidized</td>
<td align="center" valign="bottom">98 (3.60)</td>
<td align="center" valign="bottom">9 (9.18)</td>
<td align="center" valign="bottom">89 (90.82)</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Type of user</td>
</tr>
<tr>
<td align="left" valign="bottom">Policyholder</td>
<td align="center" valign="bottom">1979 (72.60)</td>
<td align="center" valign="bottom">281 (14.2)</td>
<td align="center" valign="bottom">1,698 (85.8)</td>
<td align="center" valign="middle" rowspan="2">0.014</td>
</tr>
<tr>
<td align="left" valign="bottom">Dependent</td>
<td align="center" valign="bottom">747 (27.40)</td>
<td align="center" valign="bottom">73 (9.77)</td>
<td align="center" valign="bottom">674 (90.23)</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Residence (region)</td>
</tr>
<tr>
<td align="left" valign="bottom">Bogota</td>
<td align="center" valign="bottom">1,024 (37.57)</td>
<td align="center" valign="bottom">304 (29.69)</td>
<td align="center" valign="bottom">720 (70.31)</td>
<td align="center" valign="middle" rowspan="6">
<bold>&#x003C;0.001</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">Caribbean</td>
<td align="center" valign="bottom">234 (8.58)</td>
<td align="center" valign="bottom">4 (1.71)</td>
<td align="center" valign="bottom">230 (98.29)</td>
</tr>
<tr>
<td align="left" valign="bottom">Central</td>
<td align="center" valign="bottom">645 (23.66)</td>
<td align="center" valign="bottom">5 (0.78)</td>
<td align="center" valign="bottom">640 (99.22)</td>
</tr>
<tr>
<td align="left" valign="bottom">Eastern</td>
<td align="center" valign="bottom">530 (19.44)</td>
<td align="center" valign="bottom">36 (6.79)</td>
<td align="center" valign="bottom">494 (93.21)</td>
</tr>
<tr>
<td align="left" valign="bottom">Pacific</td>
<td align="center" valign="bottom">266 (9.76)</td>
<td align="center" valign="bottom">4 (1.5)</td>
<td align="center" valign="bottom">262 (98.5)</td>
</tr>
<tr>
<td align="left" valign="bottom">Other</td>
<td align="center" valign="bottom">27 (0.99)</td>
<td align="center" valign="bottom">1 (3.7)</td>
<td align="center" valign="bottom">26 (96.3)</td>
</tr>
<tr>
<td align="left" valign="bottom">Armed conflict victim</td>
<td align="center" valign="bottom">75 (2.75)</td>
<td align="center" valign="bottom">3 (4)</td>
<td align="center" valign="bottom">72 (96)</td>
<td align="center" valign="bottom">0.3</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Socioeconomic status</td>
</tr>
<tr>
<td align="left" valign="bottom">Low income</td>
<td align="center" valign="bottom">62 (2.27)</td>
<td align="center" valign="bottom">11 (17.74)</td>
<td align="center" valign="bottom">51 (82.26)</td>
<td align="center" valign="middle" rowspan="3">0.37</td>
</tr>
<tr>
<td align="left" valign="bottom">Middle income</td>
<td align="center" valign="bottom">2,637 (86.47)</td>
<td align="center" valign="bottom">341 (12.93)</td>
<td align="center" valign="bottom">2,296 (87.07)</td>
</tr>
<tr>
<td align="left" valign="bottom">High income</td>
<td align="center" valign="bottom">27 (0.99)</td>
<td align="center" valign="bottom">2 (7.41)</td>
<td align="center" valign="bottom">25 (92.59)</td>
</tr>
<tr>
<td align="left" valign="bottom">Borderline personality disorder features</td>
<td align="center" valign="bottom">113 (4.15)</td>
<td align="center" valign="bottom">37 (32.74)</td>
<td align="center" valign="bottom">76 (67.26)</td>
<td align="center" valign="bottom">
<bold>&#x003C;0.001</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">Hypertension</td>
<td align="center" valign="bottom">815 (29.9)</td>
<td align="center" valign="bottom">78 (9.57)</td>
<td align="center" valign="bottom">737 (90.43)</td>
<td align="center" valign="bottom">
<bold>&#x003C;0.001</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">COPD</td>
<td align="center" valign="bottom">115 (4.22)</td>
<td align="center" valign="bottom">16 (13.91)</td>
<td align="center" valign="bottom">99 (86.09)</td>
<td align="center" valign="bottom">0.872</td>
</tr>
<tr>
<td align="left" valign="bottom">Hepatitis</td>
<td align="center" valign="bottom">5 (0.18)</td>
<td align="center" valign="bottom">1 (20)</td>
<td align="center" valign="bottom">4 (80)</td>
<td align="center" valign="bottom">1</td>
</tr>
<tr>
<td align="left" valign="bottom">Hypothyroidism</td>
<td align="center" valign="bottom">618 (22.67)</td>
<td align="center" valign="bottom">116 (18.77)</td>
<td align="center" valign="bottom">502 (81.23)</td>
<td align="center" valign="bottom">
<bold>&#x003C;0.001</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">Two or more medical comorbidities</td>
<td align="center" valign="bottom">246 (9.02)</td>
<td align="center" valign="bottom">36 (14.63)</td>
<td align="center" valign="bottom">210 (85.37)</td>
<td align="center" valign="bottom">0.48</td>
</tr>
</tbody>
</table>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th/>
<th align="center" valign="bottom" colspan="3">Median (Q1-Q3)</th>
<th align="center" valign="bottom"><italic>p-</italic>value</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="bottom">Age (years)</td>
<td align="center" valign="bottom">52 (18&#x2013;66)</td>
<td align="center" valign="bottom">45.5 (18&#x2013;58)</td>
<td align="center" valign="bottom">53 (18&#x2013;67)</td>
<td align="center" valign="bottom">
<bold>&#x003C;0.001</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">No. outpatient follow-ups</td>
<td align="center" valign="bottom">12 (0&#x2013;25)</td>
<td align="center" valign="bottom">3 (0&#x2013;14)</td>
<td align="center" valign="bottom">13 (0&#x2013;26)</td>
<td align="center" valign="bottom">
<bold>&#x003C;0.001</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">Psychiatric emergency visits</td>
<td align="center" valign="bottom">0 (0&#x2013;1)</td>
<td align="center" valign="bottom">1.0 (0&#x2013;2)</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">
<bold>&#x003C;0.001</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">Medical emergency visits</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0.0 (0&#x2013;1)</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0.354</td>
</tr>
<tr>
<td align="left" valign="bottom">No. hospitalizations due to medical conditions</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0.0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0.586</td>
</tr>
</tbody>
</table>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="middle" rowspan="2">Patient readmission set of variables (21 variables)</th>
<th align="center" valign="top">Total (<italic>N</italic> =&#x2009;352)</th>
<th align="center" valign="top">Readmission (<italic>n</italic> =&#x2009;165)</th>
<th align="center" valign="top">No readmission (<italic>n</italic> =&#x2009;187)</th>
<th align="center" valign="middle" rowspan="2"><italic>p-</italic>value</th>
</tr>
<tr>
<th align="center" valign="bottom" colspan="3"><italic>n</italic> (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="bottom">Female</td>
<td align="center" valign="bottom">228 (64.77)</td>
<td align="center" valign="bottom">112 (49.12)</td>
<td align="center" valign="bottom">116 (50.88)</td>
<td align="center" valign="bottom">0.301</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Marital status</td>
</tr>
<tr>
<td align="left" valign="bottom">Married</td>
<td align="center" valign="bottom">72 (20.45)</td>
<td align="center" valign="bottom">25 (34.72)</td>
<td align="center" valign="bottom">47 (65.28)</td>
<td align="center" valign="middle" rowspan="3">0.113&#x002A;</td>
</tr>
<tr>
<td align="left" valign="bottom">Divorced</td>
<td align="center" valign="bottom">7 (1.99)</td>
<td align="center" valign="bottom">4 (57.14)</td>
<td align="center" valign="bottom">3 (42.86)</td>
</tr>
<tr>
<td align="left" valign="bottom">Single</td>
<td align="center" valign="bottom">254 (72.16)</td>
<td align="center" valign="bottom">126 (49.61)</td>
<td align="center" valign="bottom">128 (50.39)</td>
</tr>
<tr>
<td align="left" valign="bottom">Common law marriage</td>
<td align="center" valign="bottom">18 (5.11)</td>
<td align="center" valign="bottom">10 (55.56)</td>
<td align="center" valign="bottom">8 (44.44)</td>
<td rowspan="2"/>
</tr>
<tr>
<td align="left" valign="bottom">Widowed</td>
<td align="center" valign="bottom">1 (0.28)</td>
<td align="center" valign="bottom">0 (0)</td>
<td align="center" valign="bottom">1 (100)</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Healthcare regime</td>
</tr>
<tr>
<td align="left" valign="bottom">Contributive</td>
<td align="center" valign="bottom">304 (86.36)</td>
<td align="center" valign="bottom">139 (45.72)</td>
<td align="center" valign="bottom">165 (54.28)</td>
<td align="center" valign="middle" rowspan="3">0.529</td>
</tr>
<tr>
<td align="left" valign="bottom">Premium plan</td>
<td align="center" valign="bottom">38 (10.80)</td>
<td align="center" valign="bottom">21 (55.26)</td>
<td align="center" valign="bottom">17 (44.74)</td>
</tr>
<tr>
<td align="left" valign="bottom">Subsidized</td>
<td align="center" valign="bottom">10 (2.84)</td>
<td align="center" valign="bottom">5 (50)</td>
<td align="center" valign="bottom">5 (50)</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Type of user</td>
</tr>
<tr>
<td align="left" valign="bottom">Policyholder</td>
<td align="center" valign="bottom">283 (80.4)</td>
<td align="center" valign="bottom">136 (48.06)</td>
<td align="center" valign="bottom">147 (51.94)</td>
<td align="center" valign="middle" rowspan="2">0.444</td>
</tr>
<tr>
<td align="left" valign="bottom">Dependent</td>
<td align="center" valign="bottom">69 (19.6)</td>
<td align="center" valign="bottom">29 (42.03)</td>
<td align="center" valign="bottom">40 (57.97)</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Residence (region)</td>
</tr>
<tr>
<td align="left" valign="bottom">Bogota</td>
<td align="center" valign="bottom">306 (86.93)</td>
<td align="center" valign="bottom">146 (47.71)</td>
<td align="center" valign="bottom">160 (52.29)</td>
<td align="center" valign="middle" rowspan="6">0.651&#x002A;</td>
</tr>
<tr>
<td align="left" valign="bottom">Caribbean</td>
<td align="center" valign="bottom">4 (1.14)</td>
<td align="center" valign="bottom">3 (75)</td>
<td align="center" valign="bottom">1 (25)</td>
</tr>
<tr>
<td align="left" valign="bottom">Central</td>
<td align="center" valign="bottom">4 (1.14)</td>
<td align="center" valign="bottom">2 (50)</td>
<td align="center" valign="bottom">2 (50)</td>
</tr>
<tr>
<td align="left" valign="bottom">Eastern</td>
<td align="center" valign="bottom">34 (9.66)</td>
<td align="center" valign="bottom">13 (38.24)</td>
<td align="center" valign="bottom">21 (61.76)</td>
</tr>
<tr>
<td align="left" valign="bottom">Pacific</td>
<td align="center" valign="bottom">3 (0.85)</td>
<td align="center" valign="bottom">1 (33.33)</td>
<td align="center" valign="bottom">2 (66.67)</td>
</tr>
<tr>
<td align="left" valign="bottom">Other</td>
<td align="center" valign="bottom">1 (0.28)</td>
<td align="center" valign="bottom">0 (0)</td>
<td align="center" valign="bottom">1 (100)</td>
</tr>
<tr>
<td align="left" valign="bottom">Armed conflict victim</td>
<td align="center" valign="bottom">3 (0.85)</td>
<td align="center" valign="bottom">2 (66.67)</td>
<td align="center" valign="bottom">1 (33.33)</td>
<td align="center" valign="bottom">0.602&#x002A;</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Socioeconomic status</td>
</tr>
<tr>
<td align="left" valign="bottom">Low income</td>
<td align="center" valign="bottom">11 (3.12)</td>
<td align="center" valign="bottom">7 (63.64)</td>
<td align="center" valign="bottom">4 (36.36)</td>
<td align="center" valign="middle" rowspan="3">0.468&#x002A;</td>
</tr>
<tr>
<td align="left" valign="bottom">Middle Income</td>
<td align="center" valign="bottom">338 (96.02)</td>
<td align="center" valign="bottom">156 (46.15)</td>
<td align="center" valign="bottom">182 (53.85)</td>
</tr>
<tr>
<td align="left" valign="bottom">High Income</td>
<td align="center" valign="bottom">3 (0.85)</td>
<td align="center" valign="bottom">2 (66.67)</td>
<td align="center" valign="bottom">1 (33.33)</td>
</tr>
<tr>
<td align="left" valign="bottom">Borderline personality disorder features</td>
<td align="center" valign="bottom">36 (10.23)</td>
<td align="center" valign="bottom">22 (61.11)</td>
<td align="center" valign="bottom">14 (38.89)</td>
<td align="center" valign="bottom">0.103</td>
</tr>
<tr>
<td align="left" valign="bottom">Hypertension</td>
<td align="center" valign="bottom">75 (21.3)</td>
<td align="center" valign="bottom">27 (36)</td>
<td align="center" valign="bottom">48 (64)</td>
<td align="center" valign="bottom">
<bold>0.046</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">COPD</td>
<td align="center" valign="bottom">14 (3.98)</td>
<td align="center" valign="bottom">5 (35.71)</td>
<td align="center" valign="bottom">9 (64.29)</td>
<td align="center" valign="bottom">0.561</td>
</tr>
<tr>
<td align="left" valign="bottom">Hepatitis</td>
<td align="center" valign="bottom">1 (0.28)</td>
<td align="center" valign="bottom">0 (0)</td>
<td align="center" valign="bottom">1 (100)</td>
<td align="center" valign="bottom">1&#x002A;</td>
</tr>
<tr>
<td align="left" valign="bottom">Hypothyroidism</td>
<td align="center" valign="bottom">119 (33.8)</td>
<td align="center" valign="bottom">59 (49.58)</td>
<td align="center" valign="bottom">60 (50.42)</td>
<td align="center" valign="bottom">0.54</td>
</tr>
<tr>
<td align="left" valign="bottom">Two or more medical comorbidities</td>
<td align="center" valign="bottom">37 (10.51)</td>
<td align="center" valign="bottom">15 (40.54)</td>
<td align="center" valign="bottom">22 (59.46)</td>
<td align="center" valign="bottom">0.521</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Place of origin (placeholder)</td>
</tr>
<tr>
<td align="left" valign="bottom">Outpatient clinic</td>
<td align="center" valign="bottom">6 (1.70)</td>
<td align="center" valign="bottom">2 (33,33)</td>
<td align="center" valign="bottom">4 (66,67)</td>
<td align="center" valign="middle" rowspan="3">0.48&#x002A;</td>
</tr>
<tr>
<td align="left" valign="bottom">Emergency department</td>
<td align="center" valign="bottom">336 (95.45)</td>
<td align="center" valign="bottom">160 (47,62)</td>
<td align="center" valign="bottom">176 (52,38)</td>
</tr>
<tr>
<td align="left" valign="bottom">Secondary transfer from other hospital</td>
<td align="center" valign="bottom">10 (2.84)</td>
<td align="center" valign="bottom">3 (30)</td>
<td align="center" valign="bottom">7 (70)</td>
</tr>
<tr>
<td align="left" valign="bottom">Admission associated with manic episode</td>
<td align="center" valign="bottom">4 (1.14)</td>
<td align="center" valign="bottom">1 (25)</td>
<td align="center" valign="bottom">3 (75)</td>
<td align="center" valign="bottom">0.706</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="5">Cause of discharge</td>
</tr>
<tr>
<td align="left" valign="bottom">Clinical improvement</td>
<td align="center" valign="bottom">326 (92.6)</td>
<td align="center" valign="bottom">152 (46,63)</td>
<td align="center" valign="bottom">174 (53,37)</td>
<td align="center" valign="middle" rowspan="5">0.91&#x002A;</td>
</tr>
<tr>
<td align="left" valign="bottom">Voluntary discharge</td>
<td align="center" valign="bottom">16 (4.55)</td>
<td align="center" valign="bottom">9 (56,25)</td>
<td align="center" valign="bottom">7 (43,75)</td>
</tr>
<tr>
<td align="left" valign="bottom">Escape</td>
<td align="center" valign="bottom">1 (0.28)</td>
<td align="center" valign="bottom">0 (0)</td>
<td align="center" valign="bottom">1 (100)</td>
</tr>
<tr>
<td align="left" valign="bottom">Home hospital</td>
<td align="center" valign="bottom">3 (0.85)</td>
<td align="center" valign="bottom">1 (33,33)</td>
<td align="center" valign="bottom">2 (66,67)</td>
</tr>
<tr>
<td align="left" valign="bottom">Transfer to another facility</td>
<td align="center" valign="bottom">6 (1.70)</td>
<td align="center" valign="bottom">3 (50)</td>
<td align="center" valign="bottom">3 (50)</td>
</tr>
</tbody>
</table>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th/>
<th align="center" valign="bottom" colspan="3">Median (Q1-Q3)</th>
<th align="center" valign="bottom"><italic>p-</italic>value</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="bottom">Age (years)</td>
<td align="center" valign="bottom">48 (36&#x2013;60)</td>
<td align="center" valign="bottom">47 (36&#x2013;57)</td>
<td align="center" valign="bottom">49 (38&#x2013;61)</td>
<td align="center" valign="bottom">0.108</td>
</tr>
<tr>
<td align="left" valign="bottom">No. outpatient follow-ups</td>
<td align="center" valign="bottom">2 (0&#x2013;7)</td>
<td align="center" valign="bottom">1 (0&#x2013;6)</td>
<td align="center" valign="bottom">3 (0&#x2013;8)</td>
<td align="center" valign="bottom">
<bold>0.021</bold>
</td>
</tr>
<tr>
<td align="left" valign="bottom">Psychiatric emergency visits</td>
<td align="center" valign="bottom">0 (0&#x2013;1)</td>
<td align="center" valign="bottom">0 (0&#x2013;1)</td>
<td align="center" valign="bottom">0 (0&#x2013;1)</td>
<td align="center" valign="bottom">0.078</td>
</tr>
<tr>
<td align="left" valign="bottom">Medical emergency visits</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0.326</td>
</tr>
<tr>
<td align="left" valign="bottom">No. hospitalizations due to medical conditions</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0 (0&#x2013;0)</td>
<td align="center" valign="bottom">0.879</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>The <italic>p-</italic>value is calculated to address the statistical significance of the difference between the readmission and no readmission groups using the Chi-squared test (&#x002A;Fisher&#x2019;s exact test) or Mann&#x2013;Whitney&#x2019;s U test. COPD, Chronic Obstructive Pulmonary Disease. Bold values are statistically significant.</p>
</table-wrap-foot>
</table-wrap>
<fig position="float" id="fig2"><label>Figure 2</label>
<caption>
<p>Study methods.</p>
</caption>
<graphic xlink:href="fpsyt-14-1266548-g002.tif"/>
</fig>
</sec>
<sec id="sec11"><label>2.4</label>
<title>Sample size determination and sampling methods</title>
<p>Sample size was calculated based on the outcome of readmission since this was deemed to be the less frequent of the two outcomes in the composite. Sample size was estimated for the outcome of readmission assuming a Cox proportional hazards model with a hazard ratio of 1.8 and an overall readmission rate of 30%. Requiring a power of 90%, a significance level &#x03B1; of 5%, and a coefficient of determination of 0.3, we obtained a final sample size of 754 subjects, split between 174 readmissions and 580 controls.</p>
<p>Since the sample was heavily skewed toward patients with no admission, we made use of oversampling techniques to balance the classes and to facilitate prediction model training. Using a Synthetic Minority Oversampling Technique (SMOTE), in which synthetic data for patients with admission were generated, an oversampling set with 1,660 patients in each group was created (<xref ref-type="bibr" rid="ref28">28</xref>). Additionally, an Edited Nearest Neighbor (ENN) method was used to generate an undersampling set which reduced the size of the larger non-admitted group, which generated a training set with 248 admitted and 1,479 non-admitted patients (<xref ref-type="bibr" rid="ref29">29</xref>). The undersampling (US), oversampling (OS), and without sampling (WS) datasets were used to train models (<xref ref-type="fig" rid="fig2">Figure 2</xref>).</p>
</sec>
<sec id="sec12"><label>2.5</label>
<title>Statistical analysis methods</title>
<p>Means and standard deviations are presented for quantitative variables, and absolute and relative frequencies are presented for categorical variables. Normality was determined using both quantile-quantile plots and the Shapiro&#x2013;Wilk test. All analyses were performed using Python software version 3.10.7 for Ubuntu version 20.04.5 LTS (Long Term Support).</p>
<sec id="sec13"><label>2.5.1</label>
<title>Logistic regression and machine learning algorithms</title>
<p>The data were randomly split into two sets in a 70:30 ratio using holdout validation, in which the training set was used to generate the prediction model for each algorithm and the test set was used to evaluate model performance. We used a more traditional statistical parametric model such as the Logistic Regression (LR) and three ML models to predict patient admission or readmission: Decision Trees (DT) (<xref ref-type="bibr" rid="ref30">30</xref>), Random Forest (RF) (<xref ref-type="bibr" rid="ref31">31</xref>), and Support Vector Machine (SVM) (<xref ref-type="bibr" rid="ref32">32</xref>) (<xref ref-type="fig" rid="fig2">Figure 2</xref>). Each of these techniques has been applied in various ways in different mental disorders, including dementia, autism spectrum disorders, and obsessive compulsive disorder (<xref ref-type="bibr" rid="ref33 ref34 ref35">33&#x2013;35</xref>). Model performance was evaluated using the area under the receiver operating characteristic curve (AUC). Additionally, we report model accuracy, precision, recall and F1 score.</p>
</sec>
<sec id="sec14"><label>2.5.2</label>
<title>Survival algorithms</title>
<p>Survival models were fitted to predict time to admission and first readmission. As with the ML models, data for survival model calibration were split into a training and test sets in a 70:30 ratio using holdout validation. Two survival models were used to estimate time to patient admission or readmission: one more standard such as the penalized Cox Model (P-Cox) (<xref ref-type="bibr" rid="ref36">36</xref>) and a ML methods such as the Random Survival Forest (RSF) (<xref ref-type="bibr" rid="ref37">37</xref>) (<xref ref-type="fig" rid="fig2">Figure 2</xref>). The performance of each model was evaluated using the concordance index (C-index), which is a generalization of the AUC which considers data censoring.</p>
</sec>
</sec>
</sec>
<sec sec-type="results" id="sec15"><label>3</label>
<title>Results</title>
<sec id="sec16"><label>3.1</label>
<title>RFE determined seven optimal variables for admission prediction</title>
<p>Regarding sex distribution, females were similarly common in either group (admission vs. no admission). Age was also a significant factor and admitted patients tended to be younger. The type of healthcare regime coverage showed a significant association with admission, with premium plan policy holders being most likely to be admitted. Patients from Bogot&#x00E1; were also more likely to be admitted than patients from the rest of Colombia. Armed conflict victim status did not differ significantly between admitted and non-admitted patients. Socioeconomic status, as systematically determined by the governmental statistical department, did not differ between groups (<xref ref-type="table" rid="tab1">Table 1</xref>).</p>
<p>Concerning prior medical history, a history of hypothyroidism or hypertension was associated with admission. Admission was more likely in patients with borderline personality disorder features and in those with prior psychiatric emergency department visits. Interestingly, patients who were not admitted tended to have a greater number of outpatient follow-up visits (<xref ref-type="table" rid="tab1">Table 1</xref>).</p>
<p>Feature engineering led to a reduced set of seven variables including psychiatric emergency visits, number of outpatient follow-ups, age, medical emergency visits, place of residence, sex, and history of hypothyroidism. Both the RFE and SFS methods were explored as alternatives for feature selection with the RFE method displaying better results. The RFE method displayed an F1 score of 0.941, compared with the F1 score of the SFS method of 0.939. Therefore, the RFE method was selected to determine the number of optimal variables (<xref ref-type="fig" rid="fig3">Figure 3A</xref>).</p>
<fig position="float" id="fig3"><label>Figure 3</label>
<caption>
<p><bold>(A)</bold> RFE method considering the outcome of admission for BD and the importance analysis for each variable. <bold>(B)</bold> SFS method considering the outcome of readmission for BD (Shaded area corresponds to the standard error).</p>
</caption>
<graphic xlink:href="fpsyt-14-1266548-g003.tif"/>
</fig>
</sec>
<sec id="sec17"><label>3.2</label>
<title>RF is the best model for patient admission prediction with an accuracy score of 0.951</title>
<p>Using the reduced set of variables to train models, the best performing model in this study was the Random Forest model, with an accuracy score of 0.951 and an AUC of 0.98 (<xref ref-type="fig" rid="fig4">Figure 4A</xref>). This best-performing model was trained using the starting dataset, without over or undersampling. The worst performing model was the Decision Tree model trained on the starting dataset (<xref ref-type="table" rid="tab2">Table 2</xref>). We obtained similar results for the survival models for time to admission, where the random survival forest obtained the best results, with a C-index of 0.95 (the P-Cox model obtained a C-index of 0.897). The median follow-up for patients in the readmission group was of 52&#x2009;months (IQR 40&#x2013;57).</p>
<fig position="float" id="fig4"><label>Figure 4</label>
<caption>
<p>ROC curve of prediction models using <bold>(A)</bold> RFE and standardized data for BD admission without sampling. <bold>(B)</bold> SFS and normalized data for BD readmission.</p>
</caption>
<graphic xlink:href="fpsyt-14-1266548-g004.tif"/>
</fig>
<table-wrap position="float" id="tab2"><label>Table 2</label>
<caption>
<p>Performance evaluation metrics for prediction of patient admission/readmission.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th/>
<th align="center" valign="bottom">Sampling type</th>
<th align="center" valign="bottom">Accuracy score</th>
<th align="center" valign="bottom">Precision Score</th>
<th align="center" valign="bottom">Recall score</th>
<th align="center" valign="bottom">F1 score</th>
<th align="center" valign="bottom">AUC</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle" colspan="7">Patient admission</td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Decision tree</td>
<td align="center" valign="middle">WS</td>
<td align="center" valign="bottom">0.93</td>
<td align="center" valign="bottom">0.70</td>
<td align="center" valign="bottom">0.75</td>
<td align="center" valign="bottom">0.72</td>
<td align="center" valign="bottom">0.85</td>
</tr>
<tr>
<td align="center" valign="middle">OS</td>
<td align="center" valign="bottom">0.92</td>
<td align="center" valign="bottom">0.64</td>
<td align="center" valign="bottom">0.83</td>
<td align="center" valign="bottom">0.72</td>
<td align="center" valign="bottom">0.88</td>
</tr>
<tr>
<td align="center" valign="middle">US</td>
<td align="center" valign="bottom">0.91</td>
<td align="center" valign="bottom">0.60</td>
<td align="center" valign="bottom">0.92</td>
<td align="center" valign="bottom">0.72</td>
<td align="center" valign="bottom">0.91</td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Random forest</td>
<td align="center" valign="middle">
<bold>WS</bold>
</td>
<td align="center" valign="bottom">
<bold>0.95</bold>
</td>
<td align="center" valign="bottom">
<bold>0.82</bold>
</td>
<td align="center" valign="bottom">
<bold>0.79</bold>
</td>
<td align="center" valign="middle">
<bold>0.81</bold>
</td>
<td align="center" valign="bottom">
<bold>0.98</bold>
</td>
</tr>
<tr>
<td align="center" valign="middle">OS</td>
<td align="center" valign="bottom">0.94</td>
<td align="center" valign="bottom">0.72</td>
<td align="center" valign="bottom">0.87</td>
<td align="center" valign="middle">0.79</td>
<td align="center" valign="bottom">0.98</td>
</tr>
<tr>
<td align="center" valign="middle">US</td>
<td align="center" valign="bottom">0.93</td>
<td align="center" valign="bottom">0.68</td>
<td align="center" valign="bottom">0.93</td>
<td align="center" valign="middle">0.78</td>
<td align="center" valign="bottom">0.98</td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Logistic regression</td>
<td align="center" valign="middle">WS</td>
<td align="center" valign="bottom">0.91</td>
<td align="center" valign="bottom">0.76</td>
<td align="center" valign="bottom">0.44</td>
<td align="center" valign="middle">0.56</td>
<td align="center" valign="bottom">0.94</td>
</tr>
<tr>
<td align="center" valign="middle">OS</td>
<td align="center" valign="bottom">0.92</td>
<td align="center" valign="bottom">0.63</td>
<td align="center" valign="bottom">0.94</td>
<td align="center" valign="middle">0.75</td>
<td align="center" valign="bottom">0.96</td>
</tr>
<tr>
<td align="center" valign="middle">US</td>
<td align="center" valign="bottom">0.93</td>
<td align="center" valign="bottom">0.70</td>
<td align="center" valign="bottom">0.80</td>
<td align="center" valign="middle">0.75</td>
<td align="center" valign="bottom">0.96</td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Support vector machine</td>
<td align="center" valign="middle">WS</td>
<td align="center" valign="bottom">0.95</td>
<td align="center" valign="bottom">0.80</td>
<td align="center" valign="bottom">0.79</td>
<td align="center" valign="middle">0.80</td>
<td align="center" valign="bottom">0.97</td>
</tr>
<tr>
<td align="center" valign="middle">OS</td>
<td align="center" valign="bottom">0.92</td>
<td align="center" valign="bottom">0.62</td>
<td align="center" valign="bottom">0.97</td>
<td align="center" valign="middle">0.76</td>
<td align="center" valign="bottom">0.98</td>
</tr>
<tr>
<td align="center" valign="middle">US</td>
<td align="center" valign="bottom">0.94</td>
<td align="center" valign="bottom">0.70</td>
<td align="center" valign="bottom">0.94</td>
<td align="center" valign="middle">0.81</td>
<td align="center" valign="bottom">0.98</td>
</tr>
<tr>
<td align="left" valign="bottom" colspan="7">Patient readmission</td>
</tr>
<tr>
<td align="left" valign="bottom">Decision tree</td>
<td align="center" valign="bottom">WS</td>
<td align="center" valign="bottom">0.55</td>
<td align="center" valign="bottom">0.52</td>
<td align="center" valign="bottom">0.62</td>
<td align="center" valign="bottom">0.56</td>
<td align="center" valign="bottom">0.52</td>
</tr>
<tr>
<td align="left" valign="bottom">Random forest</td>
<td align="center" valign="bottom">WS</td>
<td align="center" valign="bottom">0.55</td>
<td align="center" valign="bottom">0.52</td>
<td align="center" valign="bottom">0.60</td>
<td align="center" valign="bottom">0.56</td>
<td align="center" valign="bottom">0.58</td>
</tr>
<tr>
<td align="left" valign="bottom">Logistic regression</td>
<td align="center" valign="bottom">WS</td>
<td align="center" valign="bottom">0.55</td>
<td align="center" valign="bottom">0.52</td>
<td align="center" valign="bottom">0.46</td>
<td align="center" valign="bottom">0.49</td>
<td align="center" valign="bottom">0.48</td>
</tr>
<tr>
<td align="left" valign="bottom">Support vector machine</td>
<td align="center" valign="bottom">WS</td>
<td align="center" valign="bottom">0.52</td>
<td align="center" valign="bottom">0.49</td>
<td align="center" valign="bottom">0.48</td>
<td align="center" valign="bottom">0.49</td>
<td align="center" valign="bottom">0.50</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>Distribution type: without sampling (WS), oversampling (OS), undersampling (US). Best performing model in bold.</p>
</table-wrap-foot>
</table-wrap>
</sec>
<sec id="sec18"><label>3.3</label>
<title>SFS determined 11 optimal variables for readmission prediction</title>
<p>Significant associations between study variables and the outcome of readmission were identified only for a prior medical history of hypertension and for the number of outpatient follow-ups (<xref ref-type="table" rid="tab1">Table 1</xref>). Feature engineering led to a reduced set of 11 variables including psychiatric emergency visits, medical emergency visits, number of hospitalizations due to medical conditions, sex, type of user, place of residence, socioeconomic status, history of hypertension, history of Chronic Obstructive Pulmonary Disease (COPD), place of origin (placeholder), and cause of discharge. The SFS and RFE methods were considered in determining the optimal number of variables, ultimately obtaining a better performance under the F1 score metric for the SFS method (0.621 vs. 0.566) with the normalized data (<xref ref-type="fig" rid="fig3">Figure 3B</xref>).</p>
</sec>
<sec id="sec19"><label>3.4</label>
<title>Prediction of patient readmission is poor</title>
<p>Model performance for the outcome of patient readmission was poor, with the best performing model being again a RF (<xref ref-type="table" rid="tab2">Table 2</xref>). None of the models, however, reached an AUC above 0.70 (<xref ref-type="fig" rid="fig4">Figure 4B</xref>). Performance for the time to readmission survival models was likewise poor, with the RSF reaching a C-index of 0.592 (the penalized Cox model obtained a C-index of 0.497). The median follow-up time for patients in this group was of 18&#x2009;months (IQR 6&#x2013;35).</p>
</sec>
</sec>
<sec sec-type="discussion" id="sec20"><label>4</label>
<title>Discussion</title>
<p>The prediction of patient admissions in any chronic disease is of great importance to researchers, public health planners and administrators (<xref ref-type="bibr" rid="ref38">38</xref>, <xref ref-type="bibr" rid="ref39">39</xref>). Due to the varying features of different health systems, it is necessary to obtain different prediction tools for each population and health system to provide adequate risk management. To the best of our knowledge, this is the first study aiming to predict this outcome in a sample of subjects with BD in Colombia. We studied the outcomes of admission and readmission in subjects with BD in a large sample of 2,726 patients in Colombia, obtaining models with excellent predictive performance for the outcome of admission.</p>
<p>The ML models used in this study outperformed traditional statistical techniques (LR and P-Cox) in all cases with the random forest being the superior model in all cases. However, when comparing the results for the outcome of admission the random forest is only marginally superior to the LR model (AUC 0.98 vs. 0.94, respectively), which has the advantage of the latter of providing an easily interpretable model. This interpretability may be of critical importance in decision making. The difference between traditional statistical models and ML became more evident in the survival models for admission, in which the P-Cox model achieved a C-index of 0.897, while the SRF reached a C-index of 0.95. This difference between traditional statistical models and ML models was still evident in the survival models for readmission, though both models had very poor performance. This poor performance is likely due to the intrinsic difficulties associated with the prediction of survival (no readmission) compared with merely predicting the occurrence of the event within a given time frame. In addition, holdout validation was used to split the dataset corresponding to this outcome, which contained a smaller number of patients, leading to less precise estimates.</p>
<p>Random forest models are almost always superior to individual decision trees due to a variety of reasons. First, the random forest combines multiple decision trees which allows for a reduced variance and over-adjustment inherent to isolated trees. More robust and generalizable models can be derived from the averaging of several decision trees. Furthermore, the random forest models can also capture non-linear relationships and interactions between features. Lastly, random forest models can efficiently handle large datasets with high dimensionality, which makes them an adequate solution to complex problems like the prediction of BD admissions and readmissions.</p>
<p>All of the models explored for the outcome of admission had good to excellent performance, which was in contrast to the readmission models. ML models are known to require larger datasets than traditional statistical methods in order to produce better relative performance, and despite over and undersampling, the training dataset for readmission was considerably smaller (<xref ref-type="bibr" rid="ref40">40</xref>). Prospective validation for these models is also lacking and it is an area for further research. The results suggested that ML analytics has the potential to provide risk calculators to aid in predicting clinical prognosis (including patient admissions), for individual patients (<xref ref-type="bibr" rid="ref41">41</xref>).</p>
<p>Two prior studies have aimed to predict admission or readmission in BD patients. The study by Salem et al. aimed to predict readmission within 30&#x2009;days after inpatient treatment of patients with a Diagnostic and Statistical Manual of Mental Disorders (DSM)-IV diagnosis of BD using a SVM technique. Importantly, this study relied on features extracted from the Borderline Personality Questionnaire (BPQ), the scores of which were available for all subjects. The study found good discriminative ability with an area under the Receiver Operating Curve (ROC) of 0.86, concluding that borderline personality features, as measured using the BPQ, were good predictors of early readmission. The external validity of this study is limited by the availability of data on its primary predictor, which is not a standard feature of electronic health records of patients with BD (<xref ref-type="bibr" rid="ref42">42</xref>). A second study, by Edgcomb et al., also aimed to predict risk of readmission after 30&#x2009;days in patients with BD. This study took into consideration standard data from EHRs, not including any kind of standardized measurement. Using classification trees, this model achieved high accuracy with an area under the ROC curve of 0.88. An additional strength of this study was the interpretability of the model which can be easily adapted into medical thinking (<xref ref-type="bibr" rid="ref43">43</xref>).</p>
<p>In this study, the presence of borderline personality disorder features was significantly associated with admission in our first analyses. Prior work has identified these features as being a key predictor in rapid readmissions in patients with BD (<xref ref-type="bibr" rid="ref42">42</xref>). Although this variable was not significantly associated with readmission, it did meet the requirements for inclusion as a feature in the predictive models. Similarly, some medical comorbidities were both significantly associated with admission (hypertension, hypothyroidism), and met criteria for inclusion in the readmission models. This highlights their importance in considering the risk of admission/readmission in patients with BD.</p>
<p>A number of clinical variables related to the patient&#x2019;s mood, sleep quality, self-reported energy levels and medication adherence were not available to us, and they are known to be good predictors of both depressive and manic episodes in BD (<xref ref-type="bibr" rid="ref43">43</xref>, <xref ref-type="bibr" rid="ref44">44</xref>). This was probably due to the paucity of information registered in EHRs which is more likely to contain specific changes in the disease&#x2019;s symptoms and signs throughout clinical follow-up, instead of the evolution of all the clinical variables considered.</p>
<p>In obtaining the final sample, we were significantly limited by the availability of information. Despite having access to a larger cohort of patients with BD, data availability in EHRs limited their inclusion in the study. Both the number of registries and the information contained within them may have varied systematically. Physicians may be less inclined to describe clinical features in detail for patients who they deem low risk. Patients at greater risk of admission may also be those with poor adherence to outpatient follow-up, leading to fewer records from which to draw information. This may have led to selection bias in our sample.</p>
<p>Future work will concentrate on the prospective validation of the models obtained in this project, including the use of cross-validation methods (e.g., K-fold, leave one out) to allow more robust estimates of model performance by evaluating the model on different combinations of data. The limitations caused by the scarcity of data could be mitigated by EHR systems which can more intuitively allow the psychiatrist to rapidly register certain features of the mental exam, enabling more comprehensive work in mental health to be carried out. Additionally, the incorporation of longitudinal data, genetic analysis and dynamic modeling could facilitate the development of personalized treatment strategies that account for individual variations in disease progression and response to interventions in BD.</p>
</sec>
<sec sec-type="conclusions" id="sec21"><label>5</label>
<title>Conclusion</title>
<p>This study highlights the potential of utilizing ML techniques to predict hospital admission and readmission in patients with BD. The results demonstrate that ML models, particularly the Random Forest algorithm, exhibit superior predictive performance compared to traditional statistical methods. By leveraging EHRs and incorporating a range of sociodemographic and clinical variables, these models provide valuable insights into the factors influencing hospitalization in BD patients.</p>
<p>The use of ML techniques in psychiatric research, particularly in the context of BD, has the potential to deepen our understanding of the underlying mechanisms and pathophysiology of the condition. By uncovering novel associations and risk factors for patient admissions, these models contribute to the ongoing efforts to unravel the complexities of BD and guide future research directions.</p>
</sec>
<sec sec-type="data-availability" id="sec22">
<title>Data availability statement</title>
<p>The datasets presented in this article are not readily available due to patient confidentiality and data protection. The raw data supporting the conclusions of this article will be made available by the corresponding author only if this request is approved by the Research Ethics Committee of Fundaci&#x00F3;n Universitaria Sanitas, Bogot&#x00E1; D.C., Colombia, following patient protection regulations in Colombia. Requests to access the datasets should be directed to the corresponding author.</p>
</sec>
<sec sec-type="ethics-statement" id="sec23">
<title>Ethics statement</title>
<p>The studies involving humans were approved by Comit&#x00E9; de &#x00C9;tica en Investigaci&#x00F3;n Fundaci&#x00F3;n Universitaria Sanitas (CEIFUS). The studies were conducted in accordance with the local legislation and institutional requirements. Written informed consent for participation was not required from the participants or the participants&#x2019; legal guardians/next of kin in accordance with the national legislation and institutional requirements.</p>
</sec>
<sec sec-type="author-contributions" id="sec24">
<title>Author contributions</title>
<p>MP-A: Conceptualization, Data curation, Formal analysis, Funding acquisition, Investigation, Methodology, Software, Supervision, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing, Project administration. EM-M: Data curation, Formal analysis, Investigation, Methodology, Software, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. JMu: Data curation, Formal analysis, Investigation, Methodology, Software, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. RA-D: Data curation, Investigation, Methodology, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. GL-C: Conceptualization, Data curation, Investigation, Supervision, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. AC-J: Conceptualization, Data curation, Formal analysis, Supervision, Visualization, Writing &#x2013; review &#x0026; editing. JR-A: Conceptualization, Data curation, Formal analysis, Methodology, Supervision, Validation, Visualization, Writing &#x2013; review &#x0026; editing. MA-B: Conceptualization, Investigation, Methodology, Resources, Supervision, Validation, Visualization, Writing &#x2013; review &#x0026; editing. JMc: Conceptualization, Funding acquisition, Investigation, Methodology, Project administration, Resources, Supervision, Visualization, Writing &#x2013; review &#x0026; editing.</p>
</sec>
</body>
<back>
<sec sec-type="funding-information" id="sec25">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research, authorship, and/or publication of this article. Research reported in this article was supported by the National Program of Science, Technology and Innovation in Health of the Ministry of Science and Technology of the Republic of Colombia under award number 844&#x2013;2019 of the Pact to Generate New Knowledge.</p>
</sec>
<ack>
<p>The authors gratefully acknowledge Sanitas EPS and Fundaci&#x00F3;n Universitaria Sanitas for their contributions to this article. The authors also acknowledge the Ministry of Science and Technology of the Republic of Colombia for its continued financial support of science.</p>
</ack>
<sec sec-type="COI-statement" id="sec26">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
<p>The author(s) declared that they were an editorial board member of Frontiers, at the time of submission. This had no impact on the peer review process and the final decision.</p>
</sec>
<sec id="sec100" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec sec-type="supplementary-material" id="sec27">
<title>Supplementary material</title>
<p>The Supplementary material for this article can be found online at: <ext-link xlink:href="https://www.frontiersin.org/articles/10.3389/fpsyt.2023.1266548/full#supplementary-material" ext-link-type="uri">https://www.frontiersin.org/articles/10.3389/fpsyt.2023.1266548/full#supplementary-material</ext-link></p>
<supplementary-material xlink:href="Table_1.XLSX" id="SM1" mimetype="application/vnd.openxmlformats-officedocument.spreadsheetml.sheet" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="ref1"><label>1.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Merikangas</surname> <given-names>KR</given-names></name> <name><surname>Akiskal</surname> <given-names>HS</given-names></name> <name><surname>Angst</surname> <given-names>J</given-names></name> <name><surname>Greenberg</surname> <given-names>PE</given-names></name> <name><surname>Hirschfeld</surname> <given-names>RMA</given-names></name> <name><surname>Petukhova</surname> <given-names>M</given-names></name> <etal/></person-group>. <article-title>Lifetime and 12-month prevalence of bipolar Spectrum disorder in the National Comorbidity Survey Replication</article-title>. <source>Arch Gen Psychiatry</source>. (<year>2007</year>) <volume>64</volume>:<fpage>543</fpage>&#x2013;<lpage>52</lpage>. doi: <pub-id pub-id-type="doi">10.1001/archpsyc.64.5.543</pub-id>, PMID: <pub-id pub-id-type="pmid">17485606</pub-id></citation></ref>
<ref id="ref2"><label>2.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yatham</surname> <given-names>LN</given-names></name> <name><surname>Kennedy</surname> <given-names>SH</given-names></name> <name><surname>Parikh</surname> <given-names>SV</given-names></name> <name><surname>Schaffer</surname> <given-names>A</given-names></name> <name><surname>Bond</surname> <given-names>DJ</given-names></name> <name><surname>Frey</surname> <given-names>BN</given-names></name> <etal/></person-group>. <article-title>Canadian network for mood and anxiety treatments (CANMAT) and International Society for Bipolar Disorders (ISBD) 2018 guidelines for the management of patients with bipolar disorder</article-title>. <source>Bipolar Disord</source>. (<year>2018</year>) <volume>20</volume>:<fpage>97</fpage>&#x2013;<lpage>170</lpage>. doi: <pub-id pub-id-type="doi">10.1111/bdi.12609</pub-id></citation></ref>
<ref id="ref3"><label>3.</label><citation citation-type="book"><person-group person-group-type="author"><collab id="coll1">American Psychiatric Association</collab></person-group>. <source>Diagnostic and statistical manual of mental disorders</source>. <comment>DSM-5-TR.</comment> <publisher-loc>Arlington, TX</publisher-loc>: <publisher-name>American Psychiatric Association Publishing</publisher-name> (<year>2022</year>).</citation></ref>
<ref id="ref4"><label>4.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Judd</surname> <given-names>LL</given-names></name> <name><surname>Akiskal</surname> <given-names>HS</given-names></name></person-group>. <article-title>The prevalence and disability of bipolar spectrum disorders in the US population: re-analysis of the ECA database taking into account subthreshold cases</article-title>. <source>J Affect Disord</source>. (<year>2003</year>) <volume>73</volume>:<fpage>123</fpage>&#x2013;<lpage>31</lpage>. doi: <pub-id pub-id-type="doi">10.1016/S0165-0327(02)00332-4</pub-id>, PMID: <pub-id pub-id-type="pmid">12507745</pub-id></citation></ref>
<ref id="ref5"><label>5.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Reed</surname> <given-names>C</given-names></name> <name><surname>Goetz</surname> <given-names>I</given-names></name> <name><surname>Vieta</surname> <given-names>E</given-names></name> <name><surname>Bassi</surname> <given-names>M</given-names></name> <name><surname>Haro</surname> <given-names>JM</given-names></name></person-group>. <article-title>Work impairment in bipolar disorder patients &#x2013; results from a two-year observational study (EMBLEM)</article-title>. <source>Eur Psychiatry</source>. (<year>2010</year>) <volume>25</volume>:<fpage>338</fpage>&#x2013;<lpage>44</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.eurpsy.2010.01.001</pub-id>, PMID: <pub-id pub-id-type="pmid">20435449</pub-id></citation></ref>
<ref id="ref6"><label>6.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yatham</surname> <given-names>LN</given-names></name> <name><surname>Lecrubier</surname> <given-names>Y</given-names></name> <name><surname>Fieve</surname> <given-names>RR</given-names></name> <name><surname>Davis</surname> <given-names>KH</given-names></name> <name><surname>Harris</surname> <given-names>SD</given-names></name> <name><surname>Krishnan</surname> <given-names>AA</given-names></name></person-group>. <article-title>Quality of life in patients with bipolar I depression: data from 920 patients</article-title>. <source>Bipolar Disord</source>. (<year>2004</year>) <volume>6</volume>:<fpage>379</fpage>&#x2013;<lpage>85</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1399-5618.2004.00134.x</pub-id>, PMID: <pub-id pub-id-type="pmid">15383130</pub-id></citation></ref>
<ref id="ref7"><label>7.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Woods</surname> <given-names>SW</given-names></name></person-group>. <article-title>The economic burden of bipolar disease</article-title>. <source>J Clin Psychiatry</source>. (<year>2000</year>) <volume>61</volume>:<fpage>38</fpage>&#x2013;<lpage>41</lpage>.</citation></ref>
<ref id="ref8"><label>8.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Novick</surname> <given-names>DM</given-names></name> <name><surname>Swartz</surname> <given-names>HA</given-names></name> <name><surname>Frank</surname> <given-names>E</given-names></name></person-group>. <article-title>Suicide attempts in bipolar I and bipolar II disorder: a review and meta-analysis of the evidence</article-title>. <source>Bipolar Disord</source>. (<year>2010</year>) <volume>12</volume>:<fpage>1</fpage>&#x2013;<lpage>9</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1399-5618.2009.00786.x</pub-id>, PMID: <pub-id pub-id-type="pmid">20148862</pub-id></citation></ref>
<ref id="ref9"><label>9.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dome</surname> <given-names>P</given-names></name> <name><surname>Rihmer</surname> <given-names>Z</given-names></name> <name><surname>Gonda</surname> <given-names>X</given-names></name></person-group>. <article-title>Suicide risk in bipolar disorder: a brief review</article-title>. <source>Medicina</source>. (<year>2019</year>) <volume>55</volume>:<fpage>403</fpage>. doi: <pub-id pub-id-type="doi">10.3390/medicina55080403</pub-id>, PMID: <pub-id pub-id-type="pmid">31344941</pub-id></citation></ref>
<ref id="ref10"><label>10.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Volavka</surname> <given-names>J</given-names></name></person-group>. <article-title>Violence in schizophrenia and bipolar disorder</article-title>. <source>Psychiatr Danub</source>. (<year>2013</year>) <volume>25</volume>:<fpage>24</fpage>&#x2013;<lpage>33</lpage>. PMID: <pub-id pub-id-type="pmid">23470603</pub-id></citation></ref>
<ref id="ref11"><label>11.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Daglas</surname> <given-names>R</given-names></name> <name><surname>Y&#x00FC;cel</surname> <given-names>M</given-names></name> <name><surname>Cotton</surname> <given-names>S</given-names></name> <name><surname>Allott</surname> <given-names>K</given-names></name> <name><surname>Hetrick</surname> <given-names>S</given-names></name> <name><surname>Berk</surname> <given-names>M</given-names></name></person-group>. <article-title>Cognitive impairment in first-episode mania: a systematic review of the evidence in the acute and remission phases of the illness</article-title>. <source>Int J Bipolar Disord</source>. (<year>2015</year>) <volume>3</volume>:<fpage>9</fpage>. doi: <pub-id pub-id-type="doi">10.1186/s40345-015-0024-2</pub-id>, PMID: <pub-id pub-id-type="pmid">25914866</pub-id></citation></ref>
<ref id="ref12"><label>12.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bessonova</surname> <given-names>L</given-names></name> <name><surname>Ogden</surname> <given-names>K</given-names></name> <name><surname>Doane</surname> <given-names>MJ</given-names></name> <name><surname>O&#x2019;Sullivan</surname> <given-names>AK</given-names></name> <name><surname>Tohen</surname> <given-names>M</given-names></name></person-group>. <article-title>The economic burden of bipolar disorder in the United States: a systematic literature review</article-title>. <source>Clinicoecon Outcomes Res</source>. (<year>2020</year>) <volume>12</volume>:<fpage>481</fpage>&#x2013;<lpage>97</lpage>. doi: <pub-id pub-id-type="doi">10.2147/CEOR.S259338</pub-id>, PMID: <pub-id pub-id-type="pmid">32982338</pub-id></citation></ref>
<ref id="ref13"><label>13.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hirschfeld</surname> <given-names>RMA</given-names></name> <name><surname>Vornik</surname> <given-names>LA</given-names></name></person-group>. <article-title>Bipolar disorder--costs and comorbidity</article-title>. <source>Am J Manag Care</source>. (<year>2005</year>) <volume>11</volume>:<fpage>S85</fpage>&#x2013;<lpage>90</lpage>. PMID: <pub-id pub-id-type="pmid">16097719</pub-id></citation></ref>
<ref id="ref14"><label>14.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Simon</surname> <given-names>J</given-names></name> <name><surname>Pari</surname> <given-names>AAA</given-names></name> <name><surname>Wolstenholme</surname> <given-names>J</given-names></name> <name><surname>Berger</surname> <given-names>M</given-names></name> <name><surname>Goodwin</surname> <given-names>GM</given-names></name> <name><surname>Geddes</surname> <given-names>JR</given-names></name></person-group>. <article-title>The costs of bipolar disorder in the United Kingdom</article-title>. <source>Brain Behav</source>. (<year>2021</year>) <volume>11</volume>:<fpage>e2351</fpage>. doi: <pub-id pub-id-type="doi">10.1002/brb3.2351</pub-id>, PMID: <pub-id pub-id-type="pmid">34523820</pub-id></citation></ref>
<ref id="ref15"><label>15.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hong</surname> <given-names>J</given-names></name> <name><surname>Reed</surname> <given-names>C</given-names></name> <name><surname>Novick</surname> <given-names>D</given-names></name> <name><surname>Haro</surname> <given-names>JM</given-names></name> <name><surname>Windmeijer</surname> <given-names>F</given-names></name> <name><surname>Knapp</surname> <given-names>M</given-names></name></person-group>. <article-title>The cost of relapse for patients with a manic/mixed episode of bipolar disorder in the EMBLEM study</article-title>. <source>Pharmacoeconomics</source>. (<year>2010</year>) <volume>28</volume>:<fpage>555</fpage>&#x2013;<lpage>66</lpage>. doi: <pub-id pub-id-type="doi">10.2165/11535200-000000000-00000</pub-id>, PMID: <pub-id pub-id-type="pmid">20405969</pub-id></citation></ref>
<ref id="ref16"><label>16.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Perlis</surname> <given-names>RH</given-names></name> <name><surname>Ostacher</surname> <given-names>MJ</given-names></name> <name><surname>Patel</surname> <given-names>JK</given-names></name> <name><surname>Marangell</surname> <given-names>LB</given-names></name> <name><surname>Zhang</surname> <given-names>H</given-names></name> <name><surname>Wisniewski</surname> <given-names>SR</given-names></name> <etal/></person-group>. <article-title>Predictors of recurrence in bipolar disorder: primary outcomes from the systematic treatment enhancement program for bipolar disorder (STEP-BD)</article-title>. <source>Am J Psychiatr</source>. (<year>2006</year>) <volume>163</volume>:<fpage>217</fpage>&#x2013;<lpage>24</lpage>. doi: <pub-id pub-id-type="doi">10.1176/appi.ajp.163.2.217</pub-id>, PMID: <pub-id pub-id-type="pmid">16449474</pub-id></citation></ref>
<ref id="ref17"><label>17.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kansagara</surname> <given-names>D</given-names></name> <name><surname>Englander</surname> <given-names>H</given-names></name> <name><surname>Salanitro</surname> <given-names>A</given-names></name> <name><surname>Kagen</surname> <given-names>D</given-names></name> <name><surname>Theobald</surname> <given-names>C</given-names></name> <name><surname>Freeman</surname> <given-names>M</given-names></name> <etal/></person-group>. <article-title>Risk prediction models for hospital readmission</article-title>. <source>JAMA</source>. (<year>2011</year>) <volume>306</volume>:<fpage>1688</fpage>&#x2013;<lpage>98</lpage>. doi: <pub-id pub-id-type="doi">10.1001/jama.2011.1515</pub-id>, PMID: <pub-id pub-id-type="pmid">22009101</pub-id></citation></ref>
<ref id="ref18"><label>18.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhou</surname> <given-names>H</given-names></name> <name><surname>Della</surname> <given-names>PR</given-names></name> <name><surname>Roberts</surname> <given-names>P</given-names></name> <name><surname>Goh</surname> <given-names>L</given-names></name> <name><surname>Dhaliwal</surname> <given-names>SS</given-names></name></person-group>. <article-title>Utility of models to predict 28-day or 30-day unplanned hospital readmissions: an updated systematic review</article-title>. <source>BMJ Open</source>. (<year>2016</year>) <volume>6</volume>:<fpage>e011060</fpage>. doi: <pub-id pub-id-type="doi">10.1136/bmjopen-2016-011060</pub-id>, PMID: <pub-id pub-id-type="pmid">27354072</pub-id></citation></ref>
<ref id="ref19"><label>19.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Huang</surname> <given-names>Y</given-names></name> <name><surname>Talwar</surname> <given-names>A</given-names></name> <name><surname>Chatterjee</surname> <given-names>S</given-names></name> <name><surname>Aparasu</surname> <given-names>RR</given-names></name></person-group>. <article-title>Application of machine learning in predicting hospital readmissions: a scoping review of the literature</article-title>. <source>BMC Med Res Methodol</source>. (<year>2021</year>) <volume>21</volume>:<fpage>96</fpage>. doi: <pub-id pub-id-type="doi">10.1186/s12874-021-01284-z</pub-id>, PMID: <pub-id pub-id-type="pmid">33952192</pub-id></citation></ref>
<ref id="ref20"><label>20.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cusid&#x00F3;</surname> <given-names>J</given-names></name> <name><surname>Comalrena</surname> <given-names>J</given-names></name> <name><surname>Alavi</surname> <given-names>H</given-names></name> <name><surname>Llunas</surname> <given-names>L</given-names></name></person-group>. <article-title>Predicting hospital admissions to reduce crowding in the emergency departments</article-title>. <source>Appl Sci</source>. (<year>2022</year>) <volume>12</volume>:<fpage>10764</fpage>. doi: <pub-id pub-id-type="doi">10.3390/app122110764</pub-id></citation></ref>
<ref id="ref21"><label>21.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Monahan</surname> <given-names>AC</given-names></name> <name><surname>Feldman</surname> <given-names>SS</given-names></name></person-group>. <article-title>Models predicting hospital admission of adult patients utilizing prehospital data: systematic review using PROBAST and CHARMS</article-title>. <source>JMIR Med Inform</source>. (<year>2021</year>) <volume>9</volume>:<fpage>e30022</fpage>. doi: <pub-id pub-id-type="doi">10.2196/30022</pub-id>, PMID: <pub-id pub-id-type="pmid">34528893</pub-id></citation></ref>
<ref id="ref22"><label>22.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Morel</surname> <given-names>D</given-names></name> <name><surname>Yu</surname> <given-names>KC</given-names></name> <name><surname>Liu-Ferrara</surname> <given-names>A</given-names></name> <name><surname>Caceres-Suriel</surname> <given-names>AJ</given-names></name> <name><surname>Kurtz</surname> <given-names>SG</given-names></name> <name><surname>Tabak</surname> <given-names>YP</given-names></name></person-group>. <article-title>Predicting hospital readmission in patients with mental or substance use disorders: a machine learning approach</article-title>. <source>Int J Med Inform</source>. (<year>2020</year>) <volume>139</volume>:<fpage>104136</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.ijmedinf.2020.104136</pub-id></citation></ref>
<ref id="ref23"><label>23.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>G&#x00F3;ngora Alonso</surname> <given-names>S</given-names></name> <name><surname>Marques</surname> <given-names>G</given-names></name> <name><surname>Agarwal</surname> <given-names>D</given-names></name> <name><surname>De la Torre</surname> <given-names>DI</given-names></name> <name><surname>Franco-Mart&#x00ED;n</surname> <given-names>M</given-names></name></person-group>. <article-title>Comparison of machine learning algorithms in the prediction of hospitalized patients with schizophrenia</article-title>. <source>Sensors</source>. (<year>2022</year>) <volume>22</volume>:<fpage>2517</fpage>. doi: <pub-id pub-id-type="doi">10.3390/s22072517</pub-id>, PMID: <pub-id pub-id-type="pmid">35408133</pub-id></citation></ref>
<ref id="ref24"><label>24.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Betts</surname> <given-names>KS</given-names></name> <name><surname>Kisely</surname> <given-names>S</given-names></name> <name><surname>Alati</surname> <given-names>R</given-names></name></person-group>. <article-title>Predicting postpartum psychiatric admission using a machine learning approach</article-title>. <source>J Psychiatr Res</source>. (<year>2020</year>) <volume>130</volume>:<fpage>35</fpage>&#x2013;<lpage>40</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.jpsychires.2020.07.002</pub-id></citation></ref>
<ref id="ref25"><label>25.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rotenberg</surname> <given-names>L</given-names></name> <name><surname>Borges-J&#x00FA;nior</surname> <given-names>RG</given-names></name> <name><surname>Lafer</surname> <given-names>B</given-names></name> <name><surname>Salvini</surname> <given-names>R</given-names></name> <name><surname>Dias</surname> <given-names>RS</given-names></name> <name><surname>Dias</surname> <given-names>R</given-names></name> <etal/></person-group>. <article-title>Exploring machine learning to predict depressive relapses of bipolar disorder patients</article-title>. <source>J Affect Disord</source>. (<year>2021</year>) <volume>295</volume>:<fpage>681</fpage>&#x2013;<lpage>7</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.jad.2021.08.127</pub-id>, PMID: <pub-id pub-id-type="pmid">34509784</pub-id></citation></ref>
<ref id="ref26"><label>26.</label><citation citation-type="confproc"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>F</given-names></name> <name><surname>Yang</surname> <given-names>Y</given-names></name></person-group>. <article-title>Analysis of recursive feature elimination methods</article-title>. <conf-name>Proceedings of the 28th Annual International ACM SIGIR Conference on Research and Development in Information Retrieval</conf-name>. <comment>SIGIR&#x2019;05</comment>. <publisher-loc>New York, NY</publisher-loc>: <publisher-name>Association for Computing Machinery</publisher-name> (<year>2005</year>), <fpage>633</fpage>&#x2013;<lpage>634</lpage>.</citation></ref>
<ref id="ref27"><label>27.</label><citation citation-type="confproc"><person-group person-group-type="author"><name><surname>Jovic</surname> <given-names>A</given-names></name> <name><surname>Brkic</surname> <given-names>K</given-names></name> <name><surname>Bogunovic</surname> <given-names>N.</given-names></name></person-group> <article-title>A review of feature selection methods with applications</article-title> <conf-name>2015 38th International Convention on Information and Communication Technology, Electronics and Microelectronics (MIPRO)</conf-name>. <publisher-loc>Opatija</publisher-loc>: <publisher-name>IEEE</publisher-name> (<volume>2015</volume>),. <fpage>1200</fpage>&#x2013;<lpage>1205</lpage>.</citation></ref>
<ref id="ref28"><label>28.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chawla</surname> <given-names>NV</given-names></name> <name><surname>Bowyer</surname> <given-names>KW</given-names></name> <name><surname>Hall</surname> <given-names>LO</given-names></name> <name><surname>Kegelmeyer</surname> <given-names>WP</given-names></name></person-group>. <article-title>SMOTE: synthetic minority over-sampling technique</article-title>. <source>JAIR</source>. (<year>2002</year>) <volume>16</volume>:<fpage>321</fpage>&#x2013;<lpage>57</lpage>. doi: <pub-id pub-id-type="doi">10.1613/jair.953</pub-id></citation></ref>
<ref id="ref29"><label>29.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wilson</surname> <given-names>DL</given-names></name></person-group>. <article-title>Asymptotic properties of nearest neighbor rules using edited data</article-title>. <source>IEEE Trans Syst Man Cybern</source>. (<year>1972</year>) <volume>SMC-2</volume>:<fpage>408</fpage>&#x2013;<lpage>21</lpage>. doi: <pub-id pub-id-type="doi">10.1109/TSMC.1972.4309137</pub-id></citation></ref>
<ref id="ref30"><label>30.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jenhani</surname> <given-names>I</given-names></name> <name><surname>Amor</surname> <given-names>NB</given-names></name> <name><surname>Elouedi</surname> <given-names>Z</given-names></name></person-group>. <article-title>Decision trees as possibilistic classifiers</article-title>. <source>Int J Approx Reason</source>. (<year>2008</year>) <volume>48</volume>:<fpage>784</fpage>&#x2013;<lpage>807</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.ijar.2007.12.002</pub-id></citation></ref>
<ref id="ref31"><label>31.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Breiman</surname> <given-names>L</given-names></name></person-group>. <article-title>Random forests</article-title>. <source>Mach Learn</source>. (<year>2001</year>) <volume>45</volume>:<fpage>5</fpage>&#x2013;<lpage>32</lpage>. doi: <pub-id pub-id-type="doi">10.1023/A:1010933404324</pub-id></citation></ref>
<ref id="ref32"><label>32.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cortes</surname> <given-names>C</given-names></name> <name><surname>Vapnik</surname> <given-names>V</given-names></name></person-group>. <article-title>Support-vector networks</article-title>. <source>Mach Learn</source>. (<year>1995</year>) <volume>20</volume>:<fpage>273</fpage>&#x2013;<lpage>97</lpage>. doi: <pub-id pub-id-type="doi">10.1007/BF00994018</pub-id></citation></ref>
<ref id="ref33"><label>33.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Huang</surname> <given-names>F-F</given-names></name> <name><surname>Yang</surname> <given-names>X-Y</given-names></name> <name><surname>Luo</surname> <given-names>J</given-names></name> <name><surname>Yang</surname> <given-names>X-J</given-names></name> <name><surname>Meng</surname> <given-names>F-Q</given-names></name> <name><surname>Wang</surname> <given-names>P-C</given-names></name> <etal/></person-group>. <article-title>Functional and structural MRI based obsessive-compulsive disorder diagnosis using machine learning methods</article-title>. <source>BMC Psychiatry</source>. (<year>2023</year>) <volume>23</volume>:<fpage>792</fpage>. doi: <pub-id pub-id-type="doi">10.1186/s12888-023-05299-2</pub-id></citation></ref>
<ref id="ref34"><label>34.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cohen</surname> <given-names>IL</given-names></name> <name><surname>Flory</surname> <given-names>MJ</given-names></name></person-group>. <article-title>Autism Spectrum disorder decision tree subgroups predict adaptive behavior and autism severity trajectories in children with ASD</article-title>. <source>J Autism Dev Disord</source>. (<year>2019</year>) <volume>49</volume>:<fpage>1423</fpage>&#x2013;<lpage>37</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10803-018-3830-4</pub-id></citation></ref>
<ref id="ref35"><label>35.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yang</surname> <given-names>J</given-names></name> <name><surname>Sui</surname> <given-names>H</given-names></name> <name><surname>Jiao</surname> <given-names>R</given-names></name> <name><surname>Zhang</surname> <given-names>M</given-names></name> <name><surname>Zhao</surname> <given-names>X</given-names></name> <name><surname>Wang</surname> <given-names>L</given-names></name> <etal/></person-group>. <article-title>Random-Forest-algorithm-based applications of the basic characteristics and serum and imaging biomarkers to diagnose mild cognitive impairment</article-title>. <source>Curr Alzheimer Res</source>. (<year>2022</year>) <volume>19</volume>:<fpage>76</fpage>&#x2013;<lpage>83</lpage>. doi: <pub-id pub-id-type="doi">10.2174/1567205019666220128120927</pub-id>, PMID: <pub-id pub-id-type="pmid">35088670</pub-id></citation></ref>
<ref id="ref36"><label>36.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Goeman</surname> <given-names>JJ</given-names></name></person-group>. <article-title>L1 penalized estimation in the cox proportional hazards model</article-title>. <source>Biom J</source>. (<year>2010</year>) <volume>52</volume>:<fpage>70</fpage>&#x2013;<lpage>84</lpage>. doi: <pub-id pub-id-type="doi">10.1002/bimj.200900028</pub-id></citation></ref>
<ref id="ref37"><label>37.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ishwaran</surname> <given-names>H</given-names></name> <name><surname>Kogalur</surname> <given-names>UB</given-names></name> <name><surname>Blackstone</surname> <given-names>EH</given-names></name> <name><surname>Lauer</surname> <given-names>MS</given-names></name></person-group>. <article-title>Random survival forests</article-title>. <source>Ann Appl Stat</source>. (<year>2008</year>) <volume>2</volume>:<fpage>841</fpage>&#x2013;<lpage>60</lpage>. doi: <pub-id pub-id-type="doi">10.1214/08-AOAS169</pub-id></citation></ref>
<ref id="ref38"><label>38.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Symum</surname> <given-names>H</given-names></name> <name><surname>Zayas-Castro</surname> <given-names>JL</given-names></name></person-group>. <article-title>Prediction of chronic disease-related inpatient prolonged length of stay using machine learning algorithms</article-title>. <source>Healthc Inform Res</source>. (<year>2020</year>) <volume>26</volume>:<fpage>20</fpage>&#x2013;<lpage>33</lpage>. doi: <pub-id pub-id-type="doi">10.4258/hir.2020.26.1.20</pub-id>, PMID: <pub-id pub-id-type="pmid">32082697</pub-id></citation></ref>
<ref id="ref39"><label>39.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Oh</surname> <given-names>SM</given-names></name> <name><surname>Stefani</surname> <given-names>KM</given-names></name> <name><surname>Kim</surname> <given-names>HC</given-names></name></person-group>. <article-title>Development and application of chronic disease risk prediction models</article-title>. <source>Yonsei Med J</source>. (<year>2014</year>) <volume>55</volume>:<fpage>853</fpage>&#x2013;<lpage>60</lpage>. doi: <pub-id pub-id-type="doi">10.3349/ymj.2014.55.4.853</pub-id>, PMID: <pub-id pub-id-type="pmid">24954311</pub-id></citation></ref>
<ref id="ref40"><label>40.</label><citation citation-type="other"><person-group person-group-type="author"><name><surname>Cerqueira</surname> <given-names>V</given-names></name> <name><surname>Torgo</surname> <given-names>L</given-names></name> <name><surname>Soares</surname> <given-names>C</given-names></name></person-group>. <source>Machine learning vs statistical methods for time series forecasting: size matters</source>. (<year>2019</year>). <comment>Available at:</comment> <ext-link xlink:href="https://arxiv.org/abs/1909.13316" ext-link-type="uri">https://arxiv.org/abs/1909.13316</ext-link> (Accessed June 20, 2023).</citation></ref>
<ref id="ref41"><label>41.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Passos</surname> <given-names>IC</given-names></name> <name><surname>Ballester</surname> <given-names>PL</given-names></name> <name><surname>Barros</surname> <given-names>RC</given-names></name> <name><surname>Librenza-Garcia</surname> <given-names>D</given-names></name> <name><surname>Mwangi</surname> <given-names>B</given-names></name> <name><surname>Birmaher</surname> <given-names>B</given-names></name> <etal/></person-group>. <article-title>Machine learning and big data analytics in bipolar disorder: a position paper from the International Society for Bipolar Disorders big Data Task Force</article-title>. <source>Bipolar Disord</source>. (<year>2019</year>) <volume>21</volume>:<fpage>582</fpage>&#x2013;<lpage>94</lpage>. doi: <pub-id pub-id-type="doi">10.1111/bdi.12828</pub-id>, PMID: <pub-id pub-id-type="pmid">31465619</pub-id></citation></ref>
<ref id="ref42"><label>42.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Salem</surname> <given-names>H</given-names></name> <name><surname>Ruiz</surname> <given-names>A</given-names></name> <name><surname>Hernandez</surname> <given-names>S</given-names></name> <name><surname>Wahid</surname> <given-names>K</given-names></name> <name><surname>Cao</surname> <given-names>F</given-names></name> <name><surname>Karnes</surname> <given-names>B</given-names></name> <etal/></person-group>. <article-title>Borderline personality features in inpatients with bipolar disorder: impact on course and machine learning model use to predict rapid readmission</article-title>. <source>J Psychiatr Pract</source>. (<year>2019</year>) <volume>25</volume>:<fpage>279</fpage>&#x2013;<lpage>89</lpage>. doi: <pub-id pub-id-type="doi">10.1097/PRA.0000000000000392</pub-id></citation></ref>
<ref id="ref43"><label>43.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Edgcomb</surname> <given-names>J</given-names></name> <name><surname>Shaddox</surname> <given-names>T</given-names></name> <name><surname>Hellemann</surname> <given-names>G</given-names></name> <name><surname>Brooks</surname> <given-names>JO</given-names></name></person-group>. <article-title>High-risk phenotypes of early psychiatric readmission in bipolar disorder with comorbid medical illness</article-title>. <source>Psychosomatics</source>. (<year>2019</year>) <volume>60</volume>:<fpage>563</fpage>&#x2013;<lpage>73</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.psym.2019.05.002</pub-id>, PMID: <pub-id pub-id-type="pmid">31279490</pub-id></citation></ref>
<ref id="ref44"><label>44.</label><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ortiz</surname> <given-names>A</given-names></name> <name><surname>Bradler</surname> <given-names>K</given-names></name> <name><surname>Hintze</surname> <given-names>A</given-names></name></person-group>. <article-title>Episode forecasting in bipolar disorder: is energy better than mood?</article-title> <source>Bipolar Disord</source>. (<year>2018</year>) <volume>20</volume>:<fpage>470</fpage>&#x2013;<lpage>6</lpage>. doi: <pub-id pub-id-type="doi">10.1111/bdi.12603</pub-id>, PMID: <pub-id pub-id-type="pmid">29356281</pub-id></citation></ref>
</ref-list>
</back>
</article>