<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Artif. Intell.</journal-id>
<journal-title>Frontiers in Artificial Intelligence</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Artif. Intell.</abbrev-journal-title>
<issn pub-type="epub">2624-8212</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/frai.2025.1627078</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Artificial Intelligence</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Explainable AI-driven depression detection from social media using natural language processing and black box machine learning models</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name><surname>Hameed</surname> <given-names>Sidra</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/3187834/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Nauman</surname> <given-names>Muhammad</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2364087/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Akhtar</surname> <given-names>Nadeem</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Fayyaz</surname> <given-names>Muhammad A. B.</given-names></name>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x0002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/3063671/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Nawaz</surname> <given-names>Raheel</given-names></name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Faculty of Computing, The Islamia University of Bahawalpur</institution>, <addr-line>Punjab</addr-line>, <country>Pakistan</country></aff>
<aff id="aff2"><sup>2</sup><institution>Department of Information Technology, FCIT, University of the Punjab</institution>, <addr-line>Lahore</addr-line>, <country>Pakistan</country></aff>
<aff id="aff3"><sup>3</sup><institution>OTEHM, Manchester Metropolitan University</institution>, <addr-line>Manchester</addr-line>, <country>United Kingdom</country></aff>
<aff id="aff4"><sup>4</sup><institution>Pro VC-Staffordshire University</institution>, <addr-line>Staffordshire</addr-line>, <country>United Kingdom</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: Kausik Basak, JIS Institute of Advanced Studies and Research, India</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Arya Bhattacharya, Mahindra University, India</p>
<p>Saurabh Pal, University of Calcutta, India</p></fn>
<corresp id="c001">&#x0002A;Correspondence: Muhammad A. B. Fayyaz <email>m.fayyaz&#x00040;mmu.ac.uk</email></corresp>
<fn fn-type="other" id="fn001"><p>&#x02020;ORCID: Sidra Hameed <ext-link ext-link-type="uri" xlink:href="https://orcid.org/0009-0002-3975-8778">orcid.org/0009-0002-3975-8778</ext-link></p></fn></author-notes>
<pub-date pub-type="epub">
<day>11</day>
<month>09</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>8</volume>
<elocation-id>1627078</elocation-id>
<history>
<date date-type="received">
<day>14</day>
<month>05</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>18</day>
<month>08</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x000A9; 2025 Hameed, Nauman, Akhtar, Fayyaz and Nawaz.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Hameed, Nauman, Akhtar, Fayyaz and Nawaz</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<sec>
<title>Introduction</title>
<p>Mental disorders are highly prevalent in modern society, leading to substantial personal and societal burdens. Among these, depression is one of the most common, often exacerbated by socioeconomic, clinical, and individual risk factors. With the rise of social media, user-generated content offers valuable opportunities for the early detection of mental disorders through computational approaches.</p>
</sec>
<sec>
<title>Methods</title>
<p>This study explores the early detection of depression using black-box machine learning (ML) models, including Support Vector Machines (SVM), Random Forests (RF), Extreme Gradient Boosting (XGB), and Artificial Neural Networks (ANN). Advanced Natural Language Processing (NLP) techniques TF-IDF, Latent Dirichlet Allocation (LDA), N-grams, Bag of Words (BoW), and GloVe embeddings were employed to extract linguistic and semantic features. To address the interpretability limitations of black-box models, Explainable AI (XAI) methods were integrated, specifically the Local Interpretable Model-Agnostic Explanations (LIME).</p>
</sec>
<sec>
<title>Results</title>
<p>Experimental findings demonstrate that SVM achieved the highest accuracy in detecting depression from social media data, outperforming RF and other models. The application of LIME enabled granular insights into model predictions, highlighting linguistic markers strongly aligned with established psychological research.</p>
</sec>
<sec>
<title>Discussion</title>
<p>Unlike most prior studies that focus primarily on classification accuracy, this work emphasizes both predictive performance and interpretability. The integration of LIME not only enhanced transparency and interpretability but also improved the potential clinical trustworthiness of ML-based depression detection models.</p>
</sec></abstract>
<kwd-group>
<kwd>mental illness detection</kwd>
<kwd>natural language processing</kwd>
<kwd>machine learning</kwd>
<kwd>explainable artificial intelligence</kwd>
<kwd>Local Interpretable Model-Agnostic Explanations (LIME)</kwd>
</kwd-group>
<counts>
<fig-count count="5"/>
<table-count count="9"/>
<equation-count count="10"/>
<ref-count count="111"/>
<page-count count="19"/>
<word-count count="15131"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Natural Language Processing</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<title>1 Introduction</title>
<p>Mental disorders, also known as psychiatric disorders, encompass a wide range of conditions that disrupt thoughts, emotions, and behaviors, often impairing an individual&#x00027;s ability to function in daily life (<xref ref-type="bibr" rid="B74">Orabi et al., 2018</xref>; <xref ref-type="bibr" rid="B109">Zhang T. et al., 2022</xref>). These include depression, anxiety, schizophrenia, and bipolar disorder, with causes rooted in genetic, environmental, and biological factors. Among these, depression is particularly prevalent and debilitating, affecting health, relationships, and productivity. Despite the availability of treatments like psychotherapy and pharmacotherapy (<xref ref-type="bibr" rid="B51">Ive et al., 2020</xref>), depression remains a global public health concern. According to the World Health Organization, over 322 million people suffer from depression worldwide, yet many cases go undiagnosed due to stigma and limited access to mental health services, especially in low- and middle-income countries (<xref ref-type="bibr" rid="B101">Velupillai et al., 2018</xref>; <xref ref-type="bibr" rid="B104">World Health Organization, 2021</xref>).</p>
<p>To address these challenges, researchers are increasingly leveraging digital data sources&#x02014;particularly social media&#x02014;to detect mental health issues using Natural Language Processing (NLP) and Machine Learning (ML) techniques (World Health Organization, Regional Office for the Eastern Mediterranean). Social media platforms like Twitter (X), Reddit, and Facebook offer real-time insights into users&#x00027; emotions, behavior, and potential mental distress. Recent studies have demonstrated the potential of analyzing linguistic and behavioral cues to predict mood disorders and related symptoms, such as stress, self-harm, and emotional deterioration, without requiring traditional clinical assessments (<xref ref-type="bibr" rid="B108">Yazdavar et al., 2020</xref>; <xref ref-type="bibr" rid="B73">Olusegun et al., 2023</xref>; <xref ref-type="bibr" rid="B2">AbuRaed et al., 2024</xref>; <xref ref-type="bibr" rid="B18">Chancellor and De Choudhury, 2020</xref>). These digital approaches offer scalable, non-intrusive alternatives to traditional mental health diagnostics, especially in underserved populations.</p>
<p>Natural Language Processing (NLP) is a subfield of Artificial Intelligence (AI) that has facilitated various tasks in recent years, including the management and analysis of large amounts of textual data, information extraction, sentiment analysis, emotion detection, and mental health surveillance, among others (<xref ref-type="bibr" rid="B93">Steinkamp and Cook, 2021a</xref>). Feature extraction techniques in NLP are essential for transforming unstructured textual data into structured numerical representations, thereby enabling machine learning models to perform tasks such as classification, sentiment analysis, and topic modeling. Traditional methods such as Bag of Words (BoW), Term Frequency-Inverse Document Frequency (TF-IDF), and N-grams convert text into sparse feature vectors by capturing lexical patterns and word co-occurrence frequencies (<xref ref-type="bibr" rid="B84">Salton and Buckley, 1988</xref>; <xref ref-type="bibr" rid="B59">Jurafsky and Martin, 2000</xref>). While effective for basic NLP tasks, these approaches cannot capture deeper semantic relationships and contextual nuances. To address these limitations, more advanced feature extraction techniques such as Word2Vec, GloVe (<xref ref-type="bibr" rid="B76">Pennington et al., 2014</xref>), and contextual embeddings from transformer-based models like BERT (<xref ref-type="bibr" rid="B29">Devlin et al., 2019</xref>) have been developed. These methods represent words in dense vector spaces and incorporate semantic and syntactic context, significantly improving performance across a wide range of downstream NLP applications.</p>
<p>In general, users can convey their emotions through various written formats, such as posts on social media platforms, interview transcripts, and professional notes that include patient descriptions of their mental states. X (formerly known as Twitter) is most commonly known as a platform for micro-blogging because it has a straightforward user interface that enables the publication of short stories of no more than 280 characters. Tweets posted by virtually every user are available to the general public and can be retrieved using the user&#x00027;s own X API (<xref ref-type="bibr" rid="B36">Govindasamy and Palanichamy, 2021</xref>). X enables users to analyse and understand current events and trends, regardless of their geographical location. More recently, the research work focused on determining whether a person is depressed by analyzing their tweets. More precisely, Sentiment analysis can determine whether a piece of writing has been produced in a positive, negative, or neutral tone. Comments and posts from other X users can reveal whether a user is happy or sad at any given moment. Each tweet is evaluated based on its positive, negative, or neutral sentiments. The NLP system may classify tweets as either depressive or non-depressive by identifying depressive symptoms.</p>
<p>More recently, black box ML algorithms have demonstrated exceptional performance in text classification and analysis, yielding accurate and efficient results across various applications. However, their intrinsic black box nature poses notable challenges to transparency, interpretability, and trustworthiness, especially in critical domains where comprehensible decision-making processes are essential (<xref ref-type="bibr" rid="B15">Cesarini et al., 2024</xref>; <xref ref-type="bibr" rid="B52">Jahromi et al., 2024</xref>). The lack of transparency in ML and its black box nature are significant issues in its implementation in critical domains, such as healthcare (<xref ref-type="bibr" rid="B63">Khan et al., 2024</xref>; <xref ref-type="bibr" rid="B103">Weerts et al., 2019</xref>; <xref ref-type="bibr" rid="B17">Chakraborty et al., 2021</xref>; <xref ref-type="bibr" rid="B62">Kawakura et al., 2022</xref>; <xref ref-type="bibr" rid="B72">Nauman et al., 2021</xref>). Explainable Artificial Intelligence (XAI) is a subdomain of AI that aims to improve transparency by explaining the internal decision-making processes of such models (<xref ref-type="bibr" rid="B21">Chen et al., 2021</xref>).</p>
<p>One the contrary, a few known efforts to explain the black box models in literature include SHAP (<xref ref-type="bibr" rid="B20">Chelgani et al., 2023</xref>) and LIME (<xref ref-type="bibr" rid="B40">Hakkoum et al., 2020</xref>). Recent research in explainability focuses on revealing the primary features that significantly impact a model&#x00027;s decision-making process (<xref ref-type="bibr" rid="B111">Zhang Y. et al., 2022</xref>). As AI-based systems only make predictions without explaining their rationale, there is a need for mechanisms to explain and interpret their decisions. Furthermore, Local Interpretable Model-agnostic Explanations (LIME) is an explainability technique used to interpret the predictions made by machine learning models. The LIME technque approximates a complex model with a local, interpretable one around the prediction to be explained, thereby offering insights into the model&#x00027;s behavior on individual predictions (<xref ref-type="bibr" rid="B80">Ribeiro et al., 2016b</xref>). This method is particularly valuable in domains where understanding the decision-making process is crucial, such as healthcare and mental health diagnosis, as it helps build trust and provides transparency in the model&#x00027;s decisions (<xref ref-type="bibr" rid="B79">Ribeiro et al., 2016a</xref>). By applying LIME to the ML models, we can identify which features contribute most to the predictions, thus making the model&#x00027;s decisions more understandable and actionable for stakeholders.</p>
<p>Recent research on depression detection from social media platforms like Twitter and Reddit has shown that ML and deep learning models&#x02014;such as CNNs, LSTMs, and BERT&#x02014;are effective in identifying mental health indicators from user-generated content (<xref ref-type="bibr" rid="B9">Amanat et al., 2022b</xref>; <xref ref-type="bibr" rid="B56">Ji et al., 2022b</xref>). However, key gaps remain: many models operate as black boxes with limited interpretability, which is problematic in clinical contexts requiring transparency (<xref ref-type="bibr" rid="B38">Guo et al., 2023a</xref>; <xref ref-type="bibr" rid="B47">Ibrahimov and Ali, 2024</xref>). Most studies rely on a single feature representation method, overlooking the benefits of combining traditional and semantic features. Comparative evaluations across diverse models ranging from classical ML to deep neural networks are rare, limiting insight into their relative performance. Additionally, while XAI tools like LIME and SHAP are gaining traction, their integration into end-to-end depression detection systems remains limited. Finally, many datasets are weakly labeled, often based on heuristics or self-reports without clinical validation, undermining the reliability of resulting models (<xref ref-type="bibr" rid="B18">Chancellor and De Choudhury, 2020</xref>; <xref ref-type="bibr" rid="B108">Yazdavar et al., 2020</xref>).</p>
<p>This research distinguishes itself from existing literature by employing multiple feature extraction techniques and the LIME (<xref ref-type="bibr" rid="B80">Ribeiro et al., 2016b</xref>) method to elucidate the internal decision-making processes of machine learning models. This work focuses on interpreting the detection decisions made by ML models to enhance early mental disorder detection and support healthcare professionals. The research results will enable physicians to identify life-threatening diseases in their early stages, ultimately facilitating a healthier society. This work aims to leverage advanced ML techniques and XAI to facilitate the early detection of myocardial infarction through the analysis of X data. By addressing the gap in understanding how social media conversations can be leveraged for mental health insights, this research contributes novel methodologies for pre-processing data, feature extraction, and model interpretation.</p>
<p>The main contributions of this research work are as follows:</p>
<list list-type="bullet">
<list-item><p>We applied LIME explainability uniformly to 28 different feature-classifier combinations (7 feature extraction methods &#x000D7; 4 classifiers), rather than limiting interpretation to the single highest-accuracy model as in most prior studies. This research work provides a comprehensive view of how different models make predictions in the depression detection context.</p></list-item>
<list-item><p>The research findings are structured to separately rank feature extraction methods and classifiers, eliminating cross-category confusion and allowing researchers to see the independent effect of each.</p></list-item>
<list-item><p>Beyond standard tokenisation and normalization, the proposed pipeline includes slang expansion, emoticon-to-text conversion, and a mental health-specific stopword list, in particular, tailored to the noisy and abbreviated nature of depression-related social media posts.</p></list-item>
<list-item><p>The research connected linguistic patterns highlighted by LIME to known depression-related cues in psychology and linguistics literature, providing actionable insights for mental health professionals and validating black box ML model outputs beyond raw accuracy.</p></list-item>
<list-item><p>We systematically applied LIME across all model configurations, selecting it for its model-agnostic nature, suitability for short-text explanations, and computational efficiency.</p></list-item>
</list>
</sec>
<sec id="s2">
<title>2 Background</title>
<sec>
<title>2.1 Mental disorder detection</title>
<p>Modern society is plagued by a high prevalence of mental disorders, a significant source of personal and societal suffering. It is a complex, multifactorial disease influenced by several socioeconomic and clinical factors, as well as individual risk factors (<xref ref-type="bibr" rid="B97">Thornicroft et al., 2022</xref>). Depression is a typical mental condition that can affect functioning and cause suicidal thoughts or attempts (<xref ref-type="bibr" rid="B54">Jain et al., 2022</xref>). Millions of User worldwide suffer from depression each year, which is recognized as a medical condition. Persistent unhappiness or even minor stressful life events can lead to depression, illustrating the intricate relationship between mental health, NLP, ML, and AI.</p>
<p>Various computing algorithms for the automatic analysis and representation of human language are referred to as NLP (<xref ref-type="bibr" rid="B14">Cambria and White, 2014</xref>). Within Artificial Intelligence and Computer Science, the study of NLP is of utmost significance. Research into NLP employs a wide range of theoretical frameworks and methodological approaches to enable human-computer communication using natural language. NLP is an interdisciplinary field combining elements of computer science, linguistics, and mathematics, with the fundamental objective of converting human language into executable computer instructions. Natural Language Understanding and Natural Language Generation are the two basic areas of investigation in the field of NLP (<xref ref-type="bibr" rid="B61">Kang et al., 2020</xref>).</p>
</sec>
<sec>
<title>2.2 Mental disorder detection and social media</title>
<p>Social media&#x00027;s extensive use may present opportunities to lower the prevalence of undetected mental disorders. An increasing number of research projects are investigating the connection between social media and mental health. These studies attempt to determine whether there is a causal link between social media use and negative behaviors such as stress, anxiety, depression, and suicidality (<xref ref-type="bibr" rid="B37">Guntuku et al., 2017</xref>).</p>
<p>Social media networks such as X, LinkedIn, Instagram, Snapchat, and Facebook have surged in popularity, making them one of the most important sources of readily available and easily accessible information on all facets of life (<xref ref-type="bibr" rid="B50">Islam et al., 2018</xref>). Users of these platforms can express themselves freely, share their thoughts and feelings, and discuss any topic. Users suffering from mental illnesses, such as depression, may isolate themselves and avoid social engagement (<xref ref-type="bibr" rid="B44">Hemmatirad et al., 2020</xref>). However, online platforms allow users to convey their thoughts, opinions, and sentiments regarding various topics through applications such as Facebook, X, and Instagram (<xref ref-type="bibr" rid="B50">Islam et al., 2018</xref>).</p>
</sec>
<sec>
<title>2.3 Feature extraction methods</title>
<sec>
<title>2.3.1 Latent Dirichlet Allocation</title>
<p>Both NLP and ML make use of a probabilistic model known as Latent Dirichlet Allocation (LDA) for topic modeling. LDA is based on the assumption that each text contained within a corpus can be modeled as a mixture of a limited number of underlying themes and that each word contained within a document is taken from one of those subjects. This assumption guides LDA&#x00027;s operation. By utilizing the well-known LDA approach, the limit focuses on three of the seventeen Sustainable Development Goals, while simultaneously summarizing and presenting linked subtopics (<xref ref-type="bibr" rid="B6">Al Qudah et al., 2022</xref>).</p>
<p>To construct and effectively employ an LDA model, one must first ascertain the composition of the target document&#x00027;s latent themes, such as &#x003B8; and <italic>z</italic>. The following is the revised <xref ref-type="disp-formula" rid="E1">Equation 1</xref>, where &#x003B3; and &#x003D5; are the parameters of the posterior distribution of &#x003B8; and <italic>z</italic>, respectively.</p>
<disp-formula id="E1"><label>(1)</label><mml:math id="M1"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:msub><mml:mrow><mml:mi>&#x003C6;</mml:mi></mml:mrow><mml:mrow><mml:mi>n</mml:mi><mml:mi>i</mml:mi></mml:mrow></mml:msub><mml:mi>&#x0221E;</mml:mi><mml:msub><mml:mrow><mml:mi>&#x003B2;</mml:mi></mml:mrow><mml:mrow><mml:mi>i</mml:mi><mml:msub><mml:mrow><mml:mi>w</mml:mi></mml:mrow><mml:mrow><mml:mi>n</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:msub><mml:mo class="qopname">exp</mml:mo><mml:mrow><mml:mo stretchy="false">{</mml:mo><mml:mrow><mml:mtext>&#x003A8;</mml:mtext><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:msub><mml:mrow><mml:mi>&#x003B3;</mml:mi></mml:mrow><mml:mrow><mml:mi>i</mml:mi></mml:mrow></mml:msub></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo stretchy="false">}</mml:mo></mml:mrow><mml:mo>,</mml:mo><mml:msub><mml:mrow><mml:mi>&#x003B3;</mml:mi></mml:mrow><mml:mrow><mml:mi>i</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:msub><mml:mrow><mml:mi>&#x003B1;</mml:mi></mml:mrow><mml:mrow><mml:mi>i</mml:mi></mml:mrow></mml:msub><mml:mo>&#x0002B;</mml:mo><mml:mstyle displaystyle="true"><mml:munderover accentunder="false" accent="false"><mml:mrow><mml:mo>&#x02211;</mml:mo></mml:mrow><mml:mrow><mml:mi>n</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>N</mml:mi></mml:mrow></mml:munderover></mml:mstyle><mml:msub><mml:mrow><mml:mi>&#x003C6;</mml:mi></mml:mrow><mml:mrow><mml:mi>n</mml:mi><mml:mi>i</mml:mi></mml:mrow></mml:msub></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>The LDA model&#x00027;s parameters were able to be estimated after the distribution of the hidden variables had been discovered, which made the process much simpler. <italic>M</italic> represents the total number of documents, <italic>d</italic> stands for the document ID, and <italic>dni</italic> represents the greatest possible value for <italic>n</italic> that can be derived from the expectation step. <xref ref-type="disp-formula" rid="E2">Equation 2</xref> displays these three variables in their respective spots.</p>
<disp-formula id="E2"><label>(2)</label><mml:math id="M2"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>&#x003B1;</mml:mi><mml:mo>=</mml:mo><mml:mi>&#x003B1;</mml:mi><mml:mo>-</mml:mo><mml:mi>H</mml:mi><mml:msup><mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>&#x003B1;</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mo>-</mml:mo><mml:mn>1</mml:mn></mml:mrow></mml:msup><mml:mi>g</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>&#x003B1;</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow><mml:mo>,</mml:mo><mml:msub><mml:mrow><mml:mi>&#x003B2;</mml:mi></mml:mrow><mml:mrow><mml:mi>i</mml:mi><mml:mi>j</mml:mi></mml:mrow></mml:msub><mml:mi>&#x0221E;</mml:mi><mml:mstyle displaystyle="true"><mml:munderover accentunder="false" accent="false"><mml:mrow><mml:mo>&#x02211;</mml:mo></mml:mrow><mml:mrow><mml:mi>d</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>M</mml:mi></mml:mrow></mml:munderover></mml:mstyle><mml:mstyle displaystyle="true"><mml:munderover accentunder="false" accent="false"><mml:mrow><mml:mo>&#x02211;</mml:mo></mml:mrow><mml:mrow><mml:mi>n</mml:mi><mml:mo>=</mml:mo><mml:mn>1</mml:mn></mml:mrow><mml:mrow><mml:mi>N</mml:mi></mml:mrow></mml:munderover></mml:mstyle><mml:mi>&#x003C6;</mml:mi><mml:msubsup><mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>d</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mi>n</mml:mi><mml:mi>i</mml:mi></mml:mrow><mml:mrow><mml:mo>*</mml:mo></mml:mrow></mml:msubsup><mml:mi>w</mml:mi><mml:msubsup><mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>d</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mi>n</mml:mi></mml:mrow><mml:mrow><mml:mi>j</mml:mi></mml:mrow></mml:msubsup></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>In addition, the results of making inferences about hidden variables could be utilized in calculating the target document&#x00027;s generation probability value, denoted by <italic>p</italic><sub><italic>LDA</italic></sub>(<italic>x</italic>|&#x003B2;). According to LDA, a document cannot be comprehensive unless it draws from various themes because it requires drawing from a pool of ideas. In this study, these subjects were tagged and connected into thematic groups that helped distinguish between users of diabetic mobile apps who had good and negative sentiments toward them (<xref ref-type="bibr" rid="B75">Ossai and Wickramasinghe, 2023</xref>).</p>
</sec>
<sec>
<title>2.3.2 Term frequency-inverse document frequency</title>
<p>It is possible to quantify the significance or relevance of string representations (words, phrases, lemmas, and so on) by making use of the TF-IDF measure, which is utilized in the disciplines of Information Retrieval and ML that are included in a collection of documents. This can be done by comparing the document to another collection of documents. A k-best selection method and a modified version of the TF-IDF-based approach are developed as part of the feature vectorization process. Text vectorization based on modified TF-IDF, pre-trained embedding based on Google News Corpus, and a deep neural network are all components of this system (<xref ref-type="bibr" rid="B30">Dey and Das, 2023</xref>).</p>
<p>The Term Frequency (TF) of a term or word indicates the proportion of the document&#x00027;s total words that are comprised of instances of that term, as defined in <xref ref-type="disp-formula" rid="E3">Equation 3</xref>.</p>
<disp-formula id="E3"><label>(3)</label><mml:math id="M3"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>T</mml:mi><mml:mi>F</mml:mi><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mtext class="textrm" mathvariant="normal">&#x00023;&#x000A0;of&#x000A0;times&#x000A0;term&#x000A0;appears&#x000A0;in&#x000A0;doc</mml:mtext></mml:mrow><mml:mrow><mml:mtext class="textrm" mathvariant="normal">&#x00023;&#x000A0;of&#x000A0;terms&#x000A0;in&#x000A0;doc</mml:mtext></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>A term&#x00027;s Inverse Document Frequency (IDF) reveals how frequently it appears in the total number of documents in the corpus. <xref ref-type="disp-formula" rid="E4">Equation 4</xref> defines the equation for calculating the IDF. Words that do not appear in many papers (such as terms used in technical jargon, for example) are given more consideration than those used repeatedly throughout the entire work.</p>
<disp-formula id="E4"><label>(4)</label><mml:math id="M4"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>I</mml:mi><mml:mi>D</mml:mi><mml:mi>F</mml:mi><mml:mo>=</mml:mo><mml:mo class="qopname">log</mml:mo><mml:mrow><mml:mo stretchy="true">(</mml:mo><mml:mrow><mml:mfrac><mml:mrow><mml:mtext class="textrm" mathvariant="normal">&#x00023;&#x000A0;of&#x000A0;doc&#x000A0;in&#x000A0;corpus</mml:mtext></mml:mrow><mml:mrow><mml:mtext class="textrm" mathvariant="normal">&#x00023;&#x000A0;of&#x000A0;doc&#x000A0;containing&#x000A0;term&#x000A0;in&#x000A0;corpus</mml:mtext></mml:mrow></mml:mfrac></mml:mrow><mml:mo stretchy="true">)</mml:mo></mml:mrow></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>Multiplying a term&#x00027;s TF and IDF scores yields its TF-IDF, defined in <xref ref-type="disp-formula" rid="E5">Equation 5</xref>.</p>
<disp-formula id="E5"><label>(5)</label><mml:math id="M5"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>T</mml:mi><mml:mi>F</mml:mi><mml:mo>-</mml:mo><mml:mi>I</mml:mi><mml:mi>D</mml:mi><mml:mi>F</mml:mi><mml:mo>=</mml:mo><mml:mi>T</mml:mi><mml:mi>F</mml:mi><mml:mo>&#x000D7;</mml:mo><mml:mi>I</mml:mi><mml:mi>D</mml:mi><mml:mi>F</mml:mi></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>TF-IDF benefits many tasks involving natural language processing. For instance, search engines use it to determine how relevant a document is to a user&#x00027;s query. Text summarization, topic modeling, and categorization are some other applications of TF-IDF.</p>
<p>It&#x00027;s important to remember that there are several ways to determine an individual&#x00027;s IDF score. The logarithm to the base 10 is frequently used. However, a natural logarithm is used by some bookstores. To further prevent division by zero, a single can be added to the denominator in the following manner in <xref ref-type="disp-formula" rid="E6">Equation 6</xref>.</p>
<disp-formula id="E6"><label>(6)</label><mml:math id="M6"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>I</mml:mi><mml:mi>D</mml:mi><mml:mi>F</mml:mi><mml:mo>=</mml:mo><mml:mo class="qopname">log</mml:mo><mml:mrow><mml:mo stretchy="true">(</mml:mo><mml:mrow><mml:mfrac><mml:mrow><mml:mtext class="textrm" mathvariant="normal">&#x00023;&#x000A0;of&#x000A0;doc&#x000A0;in&#x000A0;corpus</mml:mtext></mml:mrow><mml:mrow><mml:mtext class="textrm" mathvariant="normal">&#x00023;&#x000A0;of&#x000A0;doc&#x000A0;contain&#x000A0;term&#x000A0;in&#x000A0;corpus</mml:mtext><mml:mo>&#x0002B;</mml:mo><mml:mn>1</mml:mn></mml:mrow></mml:mfrac></mml:mrow><mml:mo stretchy="true">)</mml:mo></mml:mrow></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>The TF-IDF technique is a widely utilized algorithm in the domain of text classification. The algorithm&#x00027;s formula is composed of two components: TF and IDF. The TF quantifies the occurrence of words within a specific class, effectively capturing their frequency in the text. In contrast, the IDF assesses the importance of a word by considering its rarity across a collection of documents, thus mitigating the influence of commonly occurring words that provide less informational value. In this study, the TF-IDF method capitalizes on the relationship between feature words and the number of texts in which appear. However, it does not account for the variation in feature word distribution across different categories, which can adversely affect classification accuracy. Despite this limitation, the TF-IDF algorithm remains a cornerstone in text classification due to its simplicity and effectiveness in various applications (<xref ref-type="bibr" rid="B106">Xiang, 2022</xref>).</p>
</sec>
<sec>
<title>2.3.3 N-grams</title>
<p>N-grams are contiguous sequences of <italic>n</italic> items extracted from a given sample of text or speech. These sequences are fundamental in various NLP applications, such as text prediction, language modeling, and information retrieval. The strength of N-grams lies in their ability to model local context within text sequences effectively, capturing the dependencies between words or characters, which is crucial for tasks like machine translation and speech recognition (<xref ref-type="bibr" rid="B70">Manning and Sch&#x000FC;tze, 1999</xref>).</p>
</sec>
<sec>
<title>2.3.4 Bag of words</title>
<p>Bag of Words (BoW) model is a fundamental method in NLP and IR, representing text data as a collection of words without considering grammar or word order (<xref ref-type="bibr" rid="B42">Harris, 1954</xref>). The BoW model involves creating a vocabulary from all unique words in a corpus and then representing each document as a vector based on the frequency of each word within the document. This simple yet powerful technique has been widely used in tasks such as document classification, sentiment analysis, and IR, as it effectively captures the presence of words in documents, which can be indicative of their content (<xref ref-type="bibr" rid="B70">Manning and Sch&#x000FC;tze, 1999</xref>). However, one of the limitations of the BoW model is that it disregards the semantics and context of words, which can lead to a loss of important information (<xref ref-type="bibr" rid="B59">Jurafsky and Martin, 2000</xref>). Despite these limitations, BoW remains a popular choice due to its simplicity and effectiveness in various applications.</p>
</sec>
<sec>
<title>2.3.5 GloVe</title>
<p>Global Vectors for Word Representation (GloVe) is an unsupervised learning algorithm for obtaining vector representations for words (<xref ref-type="bibr" rid="B76">Pennington et al., 2014</xref>). Unlike traditional count-based methods such as the BoW or TF-IDF, GloVe leverages the global statistical information of a corpus. It constructs a co-occurrence matrix of words and captures the ratios of word co-occurrences to encode semantic relationships in a lower-dimensional space. This method allows GloVe to preserve linear substructures in the vector space, enabling analogical reasoning and capturing semantic similarities between words. GloVe has shown superior performance in various natural language processing tasks, including word analogy and word similarity benchmarks, and has become a popular choice for generating word embeddings that are used in downstream tasks such as text classification, machine translation, and sentiment analysis (<xref ref-type="bibr" rid="B13">Brochier et al., 2019</xref>).</p>
<p>The GloVe model effectively bridges the gap between count-based methods and predictive models like Word2Vec by combining the strengths of both approaches. While Word2Vec captures local context through sliding windows, GloVe integrates this with global statistical information, leading to more robust and meaningful word vectors (<xref ref-type="bibr" rid="B67">Levy and Goldberg, 2015</xref>). The ability of GloVe to capture both syntactic and semantic relationships between words is further enhanced by its ability to scale efficiently across large datasets, making it ideal for tasks that require high-quality word embeddings (<xref ref-type="bibr" rid="B68">Li et al., 2018</xref>). Additionally, GloVe&#x00027;s embeddings have been shown to perform well across different languages and domains, contributing to its widespread adoption in the NLP community for applications ranging from machine translation to question-answering systems (<xref ref-type="bibr" rid="B11">Bojanowski et al., 2017</xref>).</p>
</sec>
</sec>
<sec>
<title>2.4 Prediction models</title>
<sec>
<title>2.4.1 Artificial Neural Network</title>
<p>Artificial Neural Networks (ANNs) are computational models inspired by the structure and functioning of biological neural networks. An ANN consists of layers of interconnected nodes, or neurons, that process and transmit information. These networks are typically organized in an input layer, one or more hidden layers, and an output layer (<xref ref-type="bibr" rid="B66">LeCun et al., 2015</xref>). Each neuron applies a nonlinear activation function to the weighted sum of its inputs, enabling the network to capture complex patterns in data. ANNs have been widely applied in various domains, including computer vision, natural language processing, and speech recognition, due to their ability to learn from data and generalize to unseen examples (<xref ref-type="bibr" rid="B88">Schmidhuber, 2015</xref>). One of the key advantages of ANNs is their ability to perform hierarchical feature extraction, where higher-level representations are built from lower-level features (<xref ref-type="bibr" rid="B35">Goodfellow et al., 2016</xref>). This makes ANNs particularly effective in tasks that involve high-dimensional and unstructured data.</p>
</sec>
<sec>
<title>2.4.2 Random Forest</title>
<p>Random Forest (RF) is an ensemble learning method that operates by constructing a multitude of decision trees during training and outputting the mode of the classes (classification) or mean prediction (regression) of the individual trees (<xref ref-type="bibr" rid="B12">Breiman, 2001</xref>). The core idea behind Random Forest is to combine the predictions of multiple decision trees, each trained on a random subset of the data, to improve accuracy and control overfitting. This approach reduces variance by averaging the results, making RF highly robust against noisy data and overfitting, especially in high-dimensional spaces (<xref ref-type="bibr" rid="B69">Liaw and Wiener, 2002</xref>). Moreover, Random Forest provides an intrinsic measure of feature importance, which can be valuable in interpreting the model&#x00027;s decisions (<xref ref-type="bibr" rid="B45">Ho, 1998</xref>). Due to its versatility and performance, Random Forest has been widely adopted in various fields, including bioinformatics, finance, and remote sensing.</p>
</sec>
<sec>
<title>2.4.3 Extreme Gradient Boosting</title>
<p>Extreme Gradient Boosting (XGBoost) is an advanced implementation of gradient boosting designed to enhance the performance and efficiency of machine learning models (<xref ref-type="bibr" rid="B22">Chen and Guestrin, 2016</xref>). XGBoost builds upon the principle of gradient boosting, where models are trained sequentially to correct the errors of previous models by optimizing a loss function. XGBoost introduces several innovations, including a regularization term to prevent overfitting, and efficient handling of sparse data and missing values (<xref ref-type="bibr" rid="B23">Chen et al., 2015</xref>). Furthermore, XGBoost is designed to be highly scalable, capable of running on distributed systems and handling large datasets with millions of examples (<xref ref-type="bibr" rid="B110">Zhang et al., 2017</xref>). Due to its ability to deliver high accuracy, speed, and scalability, XGBoost has become one of the most popular and widely used machine learning algorithms, particularly in competitive data science and applied machine learning.</p>
</sec>
<sec>
<title>2.4.4 Support Vector Machine</title>
<p>Support Vector Machine (SVM) is a powerful supervised learning algorithm used primarily for classification tasks, but it can also be applied to regression problems (<xref ref-type="bibr" rid="B26">Cortes and Vapnik, 1995</xref>). SVM works by finding the optimal hyperplane that maximally separates data points of different classes in a high-dimensional space. The main objective is to maximize the margin between the nearest points of different classes, known as support vectors, to the hyperplane (<xref ref-type="bibr" rid="B100">Vapnik, 1998</xref>). This approach makes SVM highly effective in high-dimensional spaces and well-suited for complex datasets where the classes are not linearly separable. To handle such cases, SVM employs the kernel trick, which implicitly maps input features into higher-dimensional spaces, enabling the algorithm to find non-linear decision boundaries (<xref ref-type="bibr" rid="B89">Sch&#x000F6;lkopf and Smola, 2002</xref>). Due to its robustness and high accuracy, SVM has been widely used in various fields, including text classification, image recognition, and bioinformatics.</p>
</sec>
</sec>
<sec>
<title>2.5 LIME</title>
<p>Local Interpretable Model-agnostic Explanations (LIME) is a popular XAI technique designed to interpret predictions made by complex, black box ML models (<xref ref-type="bibr" rid="B80">Ribeiro et al., 2016b</xref>). LIME has been applied to enhance the interpretability of models that predict mental disorders from social media data. By providing explanations for individual predictions, LIME helps in understanding which features (words, phrases, or patterns) in the text contribute most to the detection of conditions like depression or anxiety. This transparency is crucial for validating the model&#x00027;s decisions and ensuring that they align with clinical knowledge and intuition. LIME is also valuable in educational settings and research. It aids in demonstrating the internal workings of machine learning models to students and researchers. For example, in research focusing on detecting depression from X data, LIME can be used to show the significance of specific keywords or patterns, facilitating a better understanding of the model&#x00027;s behavior and improving its design and accuracy (<xref ref-type="bibr" rid="B39">Guo et al., 2023b</xref>).</p>
<p>By incorporating LIME into mental disorder detection models, researchers and practitioners can ensure that their models are not only accurate but also interpretable and trustworthy. This makes LIME a valuable tool in developing and deploying AI-based mental health diagnostics. The use of LIME enhances the transparency of machine learning models. In mental health applications, this transparency helps in gaining the trust of clinicians and patients, as they can see which features are influencing the model&#x00027;s predictions and assess whether these align with clinical expertise and evidence (<xref ref-type="bibr" rid="B64">Khoo et al., 2024</xref>; <xref ref-type="bibr" rid="B4">Akhtar et al., 2025</xref>).</p>
</sec>
</sec>
<sec id="s3">
<title>3 Literature review</title>
<p>The detection of mental disorder through social media content is garnering significant attention from the research community (<xref ref-type="bibr" rid="B86">Santos et al., 2023</xref>; <xref ref-type="bibr" rid="B8">Amanat et al., 2022a</xref>; <xref ref-type="bibr" rid="B19">Chanda et al., 2022</xref>; <xref ref-type="bibr" rid="B55">Ji et al., 2022a</xref>; <xref ref-type="bibr" rid="B41">Haque et al., 2022</xref>; <xref ref-type="bibr" rid="B102">Wani et al., 2022</xref>; <xref ref-type="bibr" rid="B81">Rizwan et al., 2022</xref>; <xref ref-type="bibr" rid="B78">Ram&#x000ED;rez-Cifuentes et al., 2021</xref>; <xref ref-type="bibr" rid="B33">Ghosh and Anwar, 2021</xref>; <xref ref-type="bibr" rid="B71">Mohammed et al., 2021</xref>). <xref ref-type="table" rid="T1">Table 1</xref> provides a summary of related studies in the literature on mental disorder detection using machine learning and NLP techniques.</p>
<table-wrap position="float" id="T1">
<label>Table 1</label>
<caption><p>Comparison table of some literature review.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>Study</bold></th>
<th valign="top" align="left"><bold>Data source</bold></th>
<th valign="top" align="left"><bold>Methods and accuracy</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B43">Helmy et al. (2024)</xref></td>
<td valign="top" align="left">English &#x00026; Arabic Tweets</td>
<td valign="top" align="left">TF-IDF, BOW Lgbm 96.3%, RF 95.7% L-svm 95.9%, Rbf-svm 20%, LR 96.4%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B86">Santos et al. (2023)</xref></td>
<td valign="top" align="left">Twitter/X</td>
<td valign="top" align="left">LIWC 58%, 67%, 56%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B3">Adarsh et al. (2023)</xref></td>
<td valign="top" align="left">Reddit</td>
<td valign="top" align="left">SVM &#x0002B; KNN 98.05%, SVM 84.92%, DT 86.16%, RF 86.64%, XGBoost 88.48%, CNN 89.42%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B60">Kabir et al. (2023)</xref></td>
<td valign="top" align="left">Twitter/X</td>
<td valign="top" align="left">SVM 51%, 51%, 51%, 54% <break/>BiLSTM 62%, 56%, 79%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B8">Amanat et al. (2022a)</xref></td>
<td valign="top" align="left">Text Tweets</td>
<td valign="top" align="left">RNN 99%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B19">Chanda et al. (2022)</xref></td>
<td valign="top" align="left">Twitter/X</td>
<td valign="top" align="left">SVM 71%, KNN 62%, RF 54%, DT 52%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B55">Ji et al. (2022a)</xref></td>
<td valign="top" align="left">Twitter/X &#x00026; Reddit</td>
<td valign="top" align="left">CNN 78%, LSTM 80%, BiLSTM 82%, RCNN 80%, SSA 81%, RN 83%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B41">Haque et al. (2022)</xref></td>
<td valign="top" align="left">Twitter/X</td>
<td valign="top" align="left">ML 93%, TN 94.0%, TN 92.5%, BiLSTM 93%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B83">Saha et al. (2022)</xref></td>
<td valign="top" align="left">Twitter/X</td>
<td valign="top" align="left">CNN 41%, RU44%, LSTM 45% Bi-GRU 41%, Bi-LSTM 41%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B102">Wani et al. (2022)</xref></td>
<td valign="top" align="left">Twitter/X, FB Youtube</td>
<td valign="top" align="left">CNN 98.15%, Word2Vec LSTM 92.19% CNN &#x0002B; LSTM 91.48%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B28">de Souza et al. (2022)</xref></td>
<td valign="top" align="left">Reddit</td>
<td valign="top" align="left">LSTM 65%, CNN 79%, Hybrid 72%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B99">Tong et al. (2022)</xref></td>
<td valign="top" align="left">TTDD, CLPsych 2015 LSVT, Statlog, Glass</td>
<td valign="top" align="left">86%, 85%, 87%, 86%, 87%, 87%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B85">Santhosh Baboo and Amirthapriya (2022)</xref></td>
<td valign="top" align="left">Twitter/X</td>
<td valign="top" align="left">RF 73%, LR 77%, SGB 72%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B109">Zhang T. et al. (2022)</xref></td>
<td valign="top" align="left">Tweets based</td>
<td valign="top" align="left">CNN 17%, RNN 36%, Transformer based methods 17%, hybrid-based methods 30%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B54">Jain et al. (2022)</xref></td>
<td valign="top" align="left">Reddit</td>
<td valign="top" align="left">NB 74.35%, SVM 77.12% LR 77.29%, RF 77.29%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B78">Ram&#x000ED;rez-Cifuentes et al. (2021)</xref></td>
<td valign="top" align="left">Reddit</td>
<td valign="top" align="left">88%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B33">Ghosh and Anwar (2021)</xref></td>
<td valign="top" align="left">Twitter/X</td>
<td valign="top" align="left">Depression score 91%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B71">Mohammed et al. (2021)</xref></td>
<td valign="top" align="left">Bangala Data</td>
<td valign="top" align="left">DT 81.56%, RF 91.64%, AB 85.12% XGB 92.80%, GNB 91.06%, MLP 87.29%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B44">Hemmatirad et al. (2020)</xref></td>
<td valign="top" align="left">Twitter/X &#x00026; Reddit</td>
<td valign="top" align="left">Twitter/X 95%, Reddit 73%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B98">Tlachac and Rundensteiner (2020)</xref></td>
<td valign="top" align="left">Twitter</td>
<td valign="top" align="left">CNN 86%, LSTM 90%, Naive Bayes 82% NN-BiLSTM with Attention model 97%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B31">Fatima et al. (2020)</xref></td>
<td valign="top" align="left">eRisk 2018</td>
<td valign="top" align="left">LR 76%, NB 67%, SVC 67%</td>
</tr>
<tr>
<td valign="top" align="left"><xref ref-type="bibr" rid="B96">Tadesse et al. (2019)</xref></td>
<td valign="top" align="left">Reddit &#x00026; Twitter/X</td>
<td valign="top" align="left">91%</td>
</tr></tbody>
</table>
</table-wrap>
<p>More recently, <xref ref-type="bibr" rid="B49">Ibrahimov et al. (2025)</xref> emphasized model transparency alongside performance, introduced Depression X, a knowledge-infused residual attention model achieving a 7% F1 improvement while providing interpretable insights. Similarly, <xref ref-type="bibr" rid="B77">Qasim et al. (2025)</xref> utilized transformer-based architectures (e.g., BERT/RoBERTa) to assess depression severity directly from social media text. In another study, <xref ref-type="bibr" rid="B32">Friedman et al. (2024)</xref> proposed EAC-Net, an emotion-aware encoder leveraging contrastive learning and self-attention, demonstrating superior recall on depression and stress detection across multiple datasets. Expanding into multimodal input, <xref ref-type="bibr" rid="B16">Cha et al. (2024)</xref> presented MOGAM, which integrates video, text, and metadata via graph-attention mechanisms, achieving 0.87 accuracy on clinically labeled users. Additionally, <xref ref-type="bibr" rid="B5">Al Asad et al. (2024)</xref> developed a BERT&#x0002B;Bi-LSTM pipeline for both English and Arabic, highlighting the importance of explainability in achieving top F1 scores across languages.</p>
<p>In another work, <xref ref-type="bibr" rid="B48">Ibrahimov et al. (2024)</xref> highlighted the critical role of XAI frameworks in making mental health AI models transparent and trustworthy. Moving beyond surveys, <xref ref-type="bibr" rid="B24">Chen and Lin (2025)</xref> developed LLM-MTD, a large-language-model based multi-task system that simultaneously classifies depression and generates medically informed explanations, achieving state-of-the-art results on the RSDD benchmark. Empirical studies, such as those by <xref ref-type="bibr" rid="B46">Hoque et al. (2025)</xref>, demonstrate the effective application of explainability tools like SHAP and LIME in real-world educational datasets. Their model attained over 91% accuracy in detecting depression in Bangladeshi university student posts, reinforcing the value of interpretability in high-risk settings.</p>
<p><xref ref-type="bibr" rid="B96">Tadesse et al. (2019)</xref> proposed an approach to identify depression-related posts on Reddit using NLP and ML techniques. Their approach highlighted the significant improvement in detection accuracy by using a combination of linguistic features and classifiers, achieving up to 91% accuracy with a Multilayer Perceptron (MLP) classifier. Their work underscores the importance of feature selection and combination in enhancing the performance of depression detection systems. Furthermore, <xref ref-type="bibr" rid="B37">Guntuku et al. (2017)</xref> explored the detection of depression through social media data, highlighting significant advances in NLP and ML that facilitate large-scale mental health screening. Despite these technological advances, the generalizability of such studies to diverse populations and alignment with established clinical criteria remains uncertain. Furthermore, recent work highlighted the pressing need to address ethical, legal, and clinical considerations, particularly concerning data ownership, privacy protection, and the integration of these methods into existing healthcare systems.</p>
<p><xref ref-type="bibr" rid="B25">Cho et al. (2019)</xref> proposed an approach to a comprehensive evaluation of the research that has been conducted on the use of ML algorithms for the diagnosis of depression, as well as recommendations for the practical uses of ML and predicted clinical remission following treatment with citalopram for twelve weeks. The dataset consisted of 1949 sad individuals who were participating in level 1 of the Sequenced Therapy Options to Relieve Depression study. In analyzing mental health using ML techniques, the primary focus is on providing a supervised learning environment for classification.</p>
<p>Another work by <xref ref-type="bibr" rid="B43">Helmy et al. (2024)</xref> explored the application of machine learning techniques to identify signs of depression in X data. It introduces manually labeled Arabic and automatically labeled English depression corpora, evaluates various pre-processing, feature extraction, and supervised classification techniques, and demonstrates the viability of machine learning for early depression detection despite recent trends favoring deep learning. This work underscores the importance of diverse, language-specific corpora and provides valuable insights into effective combinations of methodologies for predicting depression severity. The experiments demonstrated the significant impact of feature representation and resampling techniques on classifier performance, with Random Forest (<xref ref-type="bibr" rid="B12">Breiman, 2001</xref>) and RBF-SVM (<xref ref-type="bibr" rid="B89">Sch&#x000F6;lkopf and Smola, 2002</xref>; <xref ref-type="bibr" rid="B26">Cortes and Vapnik, 1995</xref>) models showing high effectiveness across different scenarios.</p>
<p><xref ref-type="bibr" rid="B86">Santos et al. (2023)</xref> proposed the use of Mixture of Experts models combined with BERT-based approaches for predicting depression and anxiety from self-reports on social media in Portuguese was proposed. Their findings indicate that models outperform traditional feature engineering methods while also allowing for more interpretable models. The finding suggests potential improvements through modifications such as attention mechanisms, hierarchical mixtures, and multi-task learning. More recently, <xref ref-type="bibr" rid="B8">Amanat et al. (2022a)</xref> investigated the application of deep learning models for detecting signs of depression in textual data sourced from social media. The proposed framework leveraged LSTM networks and RNNs to analyse and classify text, achieving an impressive accuracy of 99% in early depression detection. The findings underscore the potential of advanced machine learning techniques to enable timely and precise identification of depressive tendencies, offering valuable support for mental health interventions and early assistance strategies.</p>
<p><xref ref-type="bibr" rid="B39">Guo et al. (2023b)</xref> investigated mental health detection using text data from online forums, employing advanced machine learning techniques, including CNNs and LSTM networks. These models captured complex patterns and contextual nuances in textual data. To enhance the interpretability of these inherently black box models, the authors used the LIME technique. The LIME provided insights into the specific language patterns and features that influenced the model&#x00027;s predictions, enabling researchers to link certain textual expressions to mental health conditions. The interpretability increased the model&#x00027;s trustworthiness and supports its integration into clinical settings, where understanding decision rationale is critical for adoption and application.</p>
<p>Furthermore, <xref ref-type="bibr" rid="B58">Joyce et al. (2023)</xref> introduced the TIFU framework to enhance the trustworthiness of AI in psychiatry by focusing on transparency and interpretability. The author emphasized the importance of explainable AI, particularly through methods like LIME, to make complex models more understandable for healthcare professionals and patients, enhancing their reliability and acceptance in mental health applications. <xref ref-type="bibr" rid="B1">Abd Yusof et al. (2017)</xref> developed a computational model to identify potential causes of depression by analyzing user-generated content. This work identified prominent causes of depression and how they evolved, highlighting differences between individuals with varying levels of neuroticism. Another study, <xref ref-type="bibr" rid="B82">Sabaneh et al. (2023)</xref> integrated several advanced methodologies, including the use of ChatGPT-3 for translating Arabic text to English, QuickUMLS (<xref ref-type="bibr" rid="B92">Soldaini and Goharian, 2016</xref>) for extracting medical concepts from the translated text, and machine learning algorithms for classification. The researchers utilized a variety of classification algorithms, such as RF, SVM, and LR, with RF achieving the highest accuracy of 80.24%.</p>
<p><xref ref-type="bibr" rid="B87">Saxena et al. (2022)</xref> explored the challenge of multi-class causal categorization of mental health issues on social media, focusing on the problem of incorrect predictions due to overlapping causal explanations. Their work identified inconsistencies in causal explanations as a key reason for varying accuracy by fine-tuning classifiers and applying LIME and Integrated Gradient methods (<xref ref-type="bibr" rid="B95">Sundararajan et al., 2017</xref>; <xref ref-type="bibr" rid="B65">Kokhlikyan et al., 2020</xref>; <xref ref-type="bibr" rid="B10">Ancona et al., 2018</xref>). The proposed approach was validated on the CAMS dataset, achieving category-wise average scores of 81.29% and 0.906 using cosine similarity and word mover&#x00027;s distance, respectively. Furthermore, <xref ref-type="bibr" rid="B3">Adarsh et al. (2023)</xref> utilized LIME to enhance the explainability of their classification model. The LIME was employed to identify and highlight specific words within social media posts that significantly contribute to the classification of posts as either containing suicidal ideations or not. The LIME was used in their approach to enhance the transparency and interpretability of the depression detection model, making it easier to understand and trust the model&#x00027;s decisions, particularly in identifying critical language markers of suicidal ideation.</p>
</sec>
<sec sec-type="methods" id="s4">
<title>4 Methods</title>
<p>This research work presents a robust approach for detecting depression from X posts. The proposed approach consists of three steps: First, preprocessing techniques and feature extraction methods; second, machine learning classifiers; and third, interpretability analysis using LIME. The comprehensive methodology is depicted in <xref ref-type="fig" rid="F1">Figure 1</xref>.</p>
<fig position="float" id="F1">
<label>Figure 1</label>
<caption><p>The proposed research methodology for depression detection using NLP and XAI techniques.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-08-1627078-g0001.tif">
<alt-text>Flowchart illustrating a three-stage depression detection process using tweets. Stage 1 involves data preprocessing and feature extraction methods: LDA, N-gram, TF-IDF, BOW, and GloVe. Stage 2 includes training and testing with algorithms: ANN, XGB, RF, and SVM. Stage 3 uses Lime for result interpretation, classifying outputs as Depression or Non-Depression.</alt-text>
</graphic>
</fig>
<sec>
<title>4.1 Data collection</title>
<p>The reliability and accuracy of any proposed system are intrinsically linked to the quality and representativeness of the data collected. As such, data serves as the cornerstone of the system&#x00027;s overall effectiveness and performance (<xref ref-type="bibr" rid="B53">Jain et al., 2020</xref>). Therefore, the collection and preparation of an appropriate dataset are essential to achieve the desired objectives. The experiments conducted in this work utilized publicly available datasets hosted on Kaggle<xref ref-type="fn" rid="fn0001"><sup>1</sup></xref>. In this work, we utilize a dataset derived from posts and comments on X. The analysis focuses on key factors, such as indications of mental illnesses like depression, as reflected in user posts and interactions. Representative instances of the data set are presented in <xref ref-type="table" rid="T2">Table 2</xref> for illustrative purposes. <xref ref-type="table" rid="T3">Table 3</xref> demonstrates the words for the posts in both categories, like &#x0201C;depression&#x0201D; and &#x0201C;non-depression&#x0201D; which are topically specific.</p>
<table-wrap position="float" id="T2">
<label>Table 2</label>
<caption><p>Labeled instances from the X dataset used for classification.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>User</bold></th>
<th valign="top" align="left"><bold>Tweet</bold></th>
<th valign="top" align="left"><bold>Class</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Scotthamilton</td>
<td valign="top" align="left">Is upset that he can&#x00027;t update his Facebook by texting it... and might cry as a result School today also. Blah!</td>
<td valign="top" align="left">Non Depression</td>
</tr>
<tr>
<td valign="top" align="left">Mattycus</td>
<td valign="top" align="left">&#x00040;Kenichan I dived many times for the ball. Managed to save 50% The rest go out of bounds</td>
<td valign="top" align="left">Non Depression</td>
</tr>
<tr>
<td valign="top" align="left">BaptisteTheFool</td>
<td valign="top" align="left">Meh... Almost Lover is the exception...this track gets me depressed every time.</td>
<td valign="top" align="left">Depression</td>
</tr>
<tr>
<td valign="top" align="left">erika_strange</td>
<td valign="top" align="left">&#x00040;infidelsarecool ugh how depressing. i want to punch something.</td>
<td valign="top" align="left">Depression</td>
</tr>
<tr>
<td valign="top" align="left">ACTinglikeamama</td>
<td valign="top" align="left">&#x00040;gigdiary I know - was a little depressed that we ate so much last night there were no leftovers today</td>
<td valign="top" align="left">Depression</td>
</tr></tbody>
</table>
</table-wrap>
<table-wrap position="float" id="T3">
<label>Table 3</label>
<caption><p>Words frequently used in depressive text.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>Depression</bold></th>
<th valign="top" align="left"><bold>Non depression</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Alone, break, blame, depressed,</td>
<td valign="top" align="left">Go, days, aww, almost, holiday</td>
</tr>
<tr>
<td valign="top" align="left">Unhappy, worry, exam, rubbish</td>
<td valign="top" align="left">UK, son, YouTube, liked, Chicago</td>
</tr>
<tr>
<td valign="top" align="left">Danny, office, upset, past, reason</td>
<td valign="top" align="left">Happy, dreamy, love, Faith, Games</td>
</tr>
<tr>
<td valign="top" align="left">Needs, dead, hmmm, random, sd</td>
<td valign="top" align="left">Awsome, money, Movie, Frnds, hills</td>
</tr>
<tr>
<td valign="top" align="left">Waiting, hurt, blocked, cry, lost</td>
<td valign="top" align="left">Really, half, mad, episode, loved</td>
</tr>
<tr>
<td valign="top" align="left">Headache, summer, death, sucks</td>
<td valign="top" align="left">Lucky, cute, girls, town, visit</td>
</tr>
<tr>
<td valign="top" align="left">Miley Cyrus, job, Painfull, Massive</td>
<td valign="top" align="left">Needs, rest, excited, joy, happy</td>
</tr>
<tr>
<td valign="top" align="left">Upset, kick, dumb, Unsuccessful</td>
<td valign="top" align="left">Haha, listening, high, puppy, oooh</td>
</tr>
<tr>
<td valign="top" align="left">Disappointed, kill, Sadly, end</td>
<td valign="top" align="left">Went, ago, finished, drink, milk</td>
</tr>
<tr>
<td valign="top" align="left">STILL, feeling, busy, dark, migraine</td>
<td valign="top" align="left">DAMN, please, play, song, dance</td>
</tr></tbody>
</table>
</table-wrap>
<p>Posts were included in the dataset if they met specific inclusion criteria: they had to be written in English, contain at least five words to ensure sufficient linguistic context for natural language processing (NLP), and include depression-related keywords such as &#x0201C;depressed,&#x0201D; &#x0201C;sad,&#x0201D; or &#x0201C;alone,&#x0201D; as identified in prior research as indicators of mental distress on social media (<xref ref-type="bibr" rid="B18">Chancellor and De Choudhury, 2020</xref>). Conversely, exclusion criteria were applied to remove posts that could compromise the reliability of the model. These excluded posts containing only emojis, links, or hashtags due to their lack of semantic and syntactic depth, advertisements or spam-like content that do not represent authentic user emotions and introduce noise, and explicitly sarcastic or humorous content, which may distort model predictions due to linguistic ambiguity.</p>
<p>For this work, a total of 1,600,000 tweets were collected, representing a diverse spectrum of user experiences related to critical factors such as depression and other mental health conditions. The data collection process involved filtering tweets using keywords indicative of mental health issues, including terms such as &#x0201C;depressed,&#x0201D; &#x0201C;anxiety,&#x0201D; and other relevant expressions.</p>
</sec>
<sec>
<title>4.2 Data pre-processing</title>
<p>Before feature selection and model training, NLP techniques were applied to pre-process the collected dataset. The initial step involved cleaning X posts from the data, resulting in a substantial dataset ready for feature extraction. The data pre-processing steps are illustrated in <xref ref-type="fig" rid="F2">Figure 2</xref>.</p>
<fig position="float" id="F2">
<label>Figure 2</label>
<caption><p>Data pre-processing.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-08-1627078-g0002.tif">
<alt-text>Flowchart illustrating data preprocessing steps: starting with a dataset, followed by removing Twitter handles, links, punctuation, numbers and special characters, stop words, then tokenization, and ending with normalization.</alt-text>
</graphic>
</fig>
<p>The following pre-processing steps were implemented:</p>
<p>Tokenisation: the process of tokenization involves splitting the textual data into individual units, typically words or tokens, to facilitate further linguistic analysis. This step enables the transformation of unstructured text into a structured format suitable for feature extraction and modeling.</p>
<p>Noise removal: to enhance the quality of the dataset, noise removal techniques were applied. This included eliminating irrelevant elements such as URLs, punctuation marks, numerical values, and common stop words that do not contribute a significant semantic meaning. By refining the data set in this way, the subsequent analysis becomes more focused and meaningful.</p>
<p>Stemming: stemming techniques were employed to reduce words to their root or base forms, thereby minimizing variations of the same word (e.g., &#x0201C;running&#x0201D; and &#x0201C;ran&#x0201D; both reduced to &#x0201C;run&#x0201D;). This step helps to normalize inflected words and consolidate similar terms, leading to a more compact and informative feature space.</p>
<p>Normalization: normalization was carried out by converting all text into lowercase, ensuring uniformity across the dataset. This step prevents the algorithm from treating words with different cases (e.g., &#x0201C;Text&#x0201D; vs. &#x0201C;text&#x0201D;) as distinct entities, thereby improving the consistency and reliability of the text representation.</p>
</sec>
<sec>
<title>4.3 Feature extraction</title>
<p>After pre-processing, we employed several standard feature extraction techniques to capture the linguistic and semantic characteristics of user-generated posts. To reduce feature dimensionality while preserving document-level semantic structure, LDA was applied, modeling 70 latent topics. The TF-IDF vectors were generated to weight word importance across the corpus, facilitating the identification of salient terms. Both unigrams and bigrams were extracted using the Scikit-learn library, limiting to the top 3000 most frequent n-grams to improve contextual understanding in short texts. Additionally, the BoW representation was used as a baseline, relying on sparse word counts without considering syntactic relationships. Finally, pre-trained GloVe embeddings were incorporated to capture semantic similarity and contextual relationships between words in dense vector form. These diverse feature representations were used as inputs for training multiple machine learning models for classification.</p>
<p>The selection of feature extraction methods in this study was guided by two key considerations. First, we aimed to include a mix of traditional lexical representations (TF-IDF, N-gram, BOW), topic modeling approaches (LDA), and dense vector embeddings (GloVe) to capture both surface-level and semantic aspects of text. This diversity allows us to evaluate how LIME explanations differ when models are trained on features with fundamentally different representational properties. Second, we selected methods that are widely used in prior depression detection and sentiment analysis research, enabling meaningful comparison with existing literature and ensuring reproducibility.</p>
<p>While alternative embedding methods such as FastText, Word2Vec, or contextual embeddings like BERT are available, GloVe was chosen because it offers strong semantic representation with relatively low computational cost, making it suitable for large-scale experiments across multiple classifiers. Additionally, GloVe embeddings are static, which ensures that any interpretability differences observed using LIME are attributable to the model and feature-classifier interaction, rather than dynamic embedding variability. This controlled setting aligns with our goal of producing a consistent, explainability-focused benchmark rather than exhaustively comparing all possible embedding types.</p>
<p>GloVe was selected as the representative dense vector embedding method in our study for several reasons. First, GloVe captures global co-occurrence statistics, allowing it to encode semantic relationships between words effectively, an important factor for short-text, depression-related posts, where subtle semantic cues may indicate emotional state. Second, its extensive prior use in depression detection and sentiment analysis literature ensures comparability with existing work. Third, preliminary trials on our dataset indicated that GloVe produced slightly higher accuracy and more consistent performance across classifiers compared to FastText, whose subword-level advantages were less pronounced in our data due to the prevalence of short, informal tokens. By including GloVe alongside statistical (TF-IDF, N-gram, BOW) and probabilistic topic-modeling (LDA) methods, we aimed to evaluate LIME explainability across a diverse spectrum of feature representations.</p>
<p>In literature, TF-IDF is widely used to identify term importance in documents; it suffers from known drawbacks: high-dimensional, sparse vector representations, and an inability to capture semantic relationships such as synonymy or context, especially in short texts like tweets (<xref ref-type="bibr" rid="B107">Xu et al., 2013</xref>; <xref ref-type="bibr" rid="B57">Joshi et al., 2020</xref>). To address these issues, we limited the TF-IDF vocabulary to the top-N frequent terms (e.g., top 5,000) to reduce dimensionality and noise, following best practices in short-text feature selection. We also incorporated semantic embeddings such as GloVe to enrich the representations with contextual meaning beyond term frequency. Finally, we employed LIME to provide <italic>post-hoc</italic> interpretability, helping validate and visualize influential terms that drive classification decisions in our depression detection models.</p>
</sec>
<sec>
<title>4.4 Model training and validation</title>
<p>We divided the dataset into three distinct subsets: training, validation, and testing. The training set comprised 70% of the total data, amounting to 1,120,000 tweets, while the validation and testing sets included 15% each, consisting of 240,000 tweets for validation and 240,000 tweets for testing. This division allows us to effectively develop and fine-tune our models while ensuring an unbiased evaluation.</p>
<p>The class labels &#x0201C;depression&#x0201D; or &#x0201C;non-depression&#x0201D; were intuitively assigned based on the presence of specific depression-related keywords in the tweets, such as &#x0201C;depressed,&#x0201D; &#x0201C;sad,&#x0201D; &#x0201C;cry,&#x0201D; and &#x0201C;alone.&#x0201D; These labels were generated using a perception-based weak labeling approach, which is common in social media mental health detection studies. While this approach facilitates large-scale data collection, it may introduce noisy labels, which we mitigated through extensive pre-processing and validation using multiple classifiers. We acknowledge this limitation and highlight it in our Discussion section, suggesting the integration of expert-driven annotation in future work.</p>
<p><xref ref-type="table" rid="T4">Table 4</xref> outlines the key hyperparameters and settings used for each classifier. These values were chosen based on iterative testing on the validation set, informed by prior literature and practical experimentation.</p>
<table-wrap position="float" id="T4">
<label>Table 4</label>
<caption><p>Classifier hyperparameters used for depression detection, selected through empirical tuning and standard practices.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>Model</bold></th>
<th valign="top" align="left"><bold>Parameter</bold></th>
<th valign="top" align="left"><bold>Value/setting</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left" rowspan="6">ANN</td>
<td valign="top" align="left">Input layer</td>
<td valign="top" align="left">TF-IDF or GloVe embeddings</td>
</tr>
<tr>
<td valign="top" align="left">Hidden layers</td>
<td valign="top" align="left">Two: 64 and 32 neurons</td>
</tr>
<tr>
<td valign="top" align="left">Activation</td>
<td valign="top" align="left">ReLU (hidden), Softmax (output)</td>
</tr>
<tr>
<td valign="top" align="left">Optimizer</td>
<td valign="top" align="left">Adam (learning rate = 0.001)</td>
</tr>
<tr>
<td valign="top" align="left">Loss function</td>
<td valign="top" align="left">Categorical Cross-entropy</td>
</tr>
<tr>
<td valign="top" align="left">Epochs/batch size</td>
<td valign="top" align="left">20/32 with early Stopping</td>
</tr>
<tr>
<td valign="top" align="left" rowspan="3">SVM</td>
<td valign="top" align="left">Kernel</td>
<td valign="top" align="left">Linear</td>
</tr>
<tr>
<td valign="top" align="left">Regularization (C)</td>
<td valign="top" align="left">1.0</td>
</tr>
<tr>
<td valign="top" align="left">Decision function</td>
<td valign="top" align="left">One-vs.-rest</td>
</tr>
<tr>
<td valign="top" align="left" rowspan="4">Random forest</td>
<td valign="top" align="left">Number of trees</td>
<td valign="top" align="left">100</td>
</tr>
<tr>
<td valign="top" align="left">Max depth</td>
<td valign="top" align="left">None (expand until pure)</td>
</tr>
<tr>
<td valign="top" align="left">Criterion</td>
<td valign="top" align="left">Gini Index</td>
</tr>
<tr>
<td valign="top" align="left">Bootstrap</td>
<td valign="top" align="left">True</td>
</tr>
<tr>
<td valign="top" align="left" rowspan="4">XGBoost</td>
<td valign="top" align="left">Number of estimators</td>
<td valign="top" align="left">100</td>
</tr>
<tr>
<td valign="top" align="left">Max depth</td>
<td valign="top" align="left">6</td>
</tr>
<tr>
<td valign="top" align="left">Learning rate</td>
<td valign="top" align="left">0.1</td>
</tr>
<tr>
<td valign="top" align="left">Objective</td>
<td valign="top" align="left">Binary:logistic</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>To check how well our models worked, we used a common method called hold-out validation. We split the dataset into 80% for training and 20% for testing. The test data was not used during training or tuning, so we could see how well the model performs on new, unseen data. We also used five-fold cross-validation while training models like SVM, Random Forest, and ANN. This method helps us fine-tune model settings and reduce the chance of overfitting. To measure performance, we used standard metrics like accuracy, precision, recall, and F1-score. We also created confusion matrices and ROC curves to visualize how the models performed. These results were all based on the test data to make sure they were reliable. For the ANN, we used early stopping to avoid overfitting. This means the training stopped automatically when the model stopped improving on the validation data.</p>
<p>To classify tweets as indicative of depression or not, we selected and implemented several well-established machine learning models that have demonstrated strong performance in text classification tasks. We utilized the ANN, specifically a Multi-Layer Perceptron (MLP) architecture with two hidden layers of 4 and 16 neurons, respectively, as this provided a manageable level of complexity for evaluating various feature representations. The RF algorithm was selected for its robustness against noisy data and its ability to mitigate overfitting through the aggregation of multiple decision trees and random feature selection. We also employed XGBoost, chosen for its superior accuracy and regularization capabilities, which are achieved by sequentially optimizing weak learners to minimize classification error. Additionally, we incorporated the SVM because of its effectiveness in handling high-dimensional feature spaces, utilizing kernel methods to identify optimal separating hyperplanes. Each of these classifiers was trained using the same preprocessed and vectorized data, ensuring a fair and consistent comparison of their respective performance in detecting depression from tweets.</p>
</sec>
<sec>
<title>4.5 Performance evaluations</title>
<p>We assessed the performance of the model by applying well-known performance metrics, including accuracy and precision, and recall. The formulas of these evaluation metrics are shown below:</p>
<disp-formula id="E7"><label>(7)</label><mml:math id="M7"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mtext class="textrm" mathvariant="normal">Accuracy</mml:mtext><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>T</mml:mi><mml:mi>N</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>T</mml:mi><mml:mi>N</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>N</mml:mi></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E8"><label>(8)</label><mml:math id="M8"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mtext class="textrm" mathvariant="normal">Precision</mml:mtext><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>P</mml:mi></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E9"><label>(9)</label><mml:math id="M9"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mtext class="textrm" mathvariant="normal">Recall</mml:mtext><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>&#x0002B;</mml:mo><mml:mi>F</mml:mi><mml:mi>N</mml:mi></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<disp-formula id="E10"><label>(10)</label><mml:math id="M10"><mml:mtable class="eqnarray" columnalign="left"><mml:mtr><mml:mtd><mml:mi>F</mml:mi><mml:mn>1</mml:mn><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mn>2</mml:mn><mml:mo>&#x000D7;</mml:mo><mml:mtext class="textrm" mathvariant="normal">Precision</mml:mtext><mml:mo>&#x000D7;</mml:mo><mml:mtext class="textrm" mathvariant="normal">Recall</mml:mtext></mml:mrow><mml:mrow><mml:mtext class="textrm" mathvariant="normal">Precision</mml:mtext><mml:mo>&#x0002B;</mml:mo><mml:mtext class="textrm" mathvariant="normal">Recall</mml:mtext></mml:mrow></mml:mfrac></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
</sec>
</sec>
<sec sec-type="results" id="s5">
<title>5 Results</title>
<p>The robustnsess of the proposed approach lies in our preprocessing pipeline, which refers to its ability to effectively clean and normalize noisy, informal social media text, including irregular spellings, emoticons, repeated characters, and abbreviations which are common on platforms like X (<xref ref-type="bibr" rid="B18">Chancellor and De Choudhury, 2020</xref>; <xref ref-type="bibr" rid="B94">Steinkamp and Cook, 2021b</xref>). This pipeline improved the quality of feature extraction by reducing vocabulary sparsity and enhancing model generalizability. Our approach also ensured consistent performance across all tested classifiers (ANN, SVM, XGBoost, RF), demonstrating resilience against data variation, which is crucial when working with user-generated, unstructured data. Compared to existing methods (e.g., <xref ref-type="bibr" rid="B55">Ji et al., 2022a</xref>; <xref ref-type="bibr" rid="B8">Amanat et al., 2022a</xref>), our framework achieves higher or comparable accuracy using simpler architectures (e.g., GloVe&#x0002B;RF: 88%, SVM&#x0002B;TFIDF: 79%), while maintaining interpretability through LIME, which is rarely integrated in similar works (<xref ref-type="bibr" rid="B56">Ji et al., 2022b</xref>; <xref ref-type="bibr" rid="B9">Amanat et al., 2022b</xref>). This hybridization of multiple NLP features with black box models, accompanied by transparent explanations, offers a practical and explainable solution for early depression detection. Furthermore, our system does not require deep learning or transformer-based models, making it more computationally efficient and suitable for real-world deployment.</p>
<p>Once the performance of the ML model has been evaluated, it becomes essential to analyse and interpret the findings to gain deeper insights into the model&#x00027;s performance. This involves discerning the crucial features influencing the model&#x00027;s predictions, comprehending the relationships between these features and the target variable, and identifying any pertinent patterns or trends within the dataset. This work employed a comprehensive set of experiments to detect depression from X posts using various combinations of feature extraction methods and ML classifiers. To visualize the most prominent terms that express emotions, a word cloud representation is utilized. <xref ref-type="fig" rid="F3">Figure 3a</xref> demonstrates the depressive user&#x00027;s feelings, experiences, and stories. This method summarizes the ideas and phrases most commonly linked to mental disorders in the analysis of social media discussions. <xref ref-type="fig" rid="F3">Figure 3b</xref> illustrates the word cloud representing the positive core words of the dataset. Word clouds are visual representations that draw attention to the most prevalent terms in a collection of text. In addition, the prominence of a word in the cloud reflects its frequency of use in the corresponding tweets. This method summarizes the ideas and phrases most commonly linked to mental disorders in the analysis of social media discussions. On the contrary, <xref ref-type="table" rid="T5">Table 5</xref> presents the words correlated with the specific topics generated from the posts. These topics comprise a lexicon of words commonly used among accounts associated with depression.</p>
<fig position="float" id="F3">
<label>Figure 3</label>
<caption><p>Word cloud for training data to visualize important. <bold>(a)</bold> Negative and <bold>(b)</bold> positive words.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-08-1627078-g0003.tif">
<alt-text>Two word clouds labeled &#x0201C;Negative Words&#x0201D; and &#x0201C;Positive Words.&#x0201D; On the left, &#x0201C;Negative Words&#x0201D; features prominent terms like &#x0201C;sleep,&#x0201D; &#x0201C;miss,&#x0201D; &#x0201C;time,&#x0201D; and &#x0201C;know&#x0201D; in various sizes. On the right, &#x0201C;Positive Words&#x0201D; includes &#x0201C;good,&#x0201D; &#x0201C;love,&#x0201D; &#x0201C;time,&#x0201D; &#x0201C;thank,&#x0201D; and &#x0201C;today&#x0201D; prominently displayed. Both clouds use a mix of purple, yellow, and green colors for the words.</alt-text>
</graphic>
</fig>
<table-wrap position="float" id="T5">
<label>Table 5</label>
<caption><p>Topics extracted with LDA.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>Sr</bold></th>
<th valign="top" align="left"><bold>Topics</bold></th>
<th valign="top" align="left"><bold>Words</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">1</td>
<td valign="top" align="left">Daily activities</td>
<td valign="top" align="left">Haha, play, friend, year, stop haha play friend year stop yeah tell think today chang want left know word</td>
</tr>
<tr>
<td valign="top" align="left">2</td>
<td valign="top" align="left">Planning:</td>
<td valign="top" align="left">Hous, damn, love, plan, like hous damn love plan like trip time dear usual watch lost today tomorrow cook list want</td>
</tr>
<tr>
<td valign="top" align="left">3</td>
<td valign="top" align="left">SocialMedia usage</td>
<td valign="top" align="left">Gonna, enjoy, thing, fun, welcome gonna enjoy thing fun welcom weather X love come cool sure</td>
</tr>
<tr>
<td valign="top" align="left">4</td>
<td valign="top" align="left">Home life</td>
<td valign="top" align="left">Home, week, away, head, update home week away head updat night guess today stuck bought outside</td>
</tr>
<tr>
<td valign="top" align="left">5</td>
<td valign="top" align="left">Depression</td>
<td valign="top" align="left">Good, feel, morn, better, hope good feel morn better hope morning right like hate night realli make today coffe bed think sleep class</td>
</tr>
<tr>
<td valign="top" align="left">6</td>
<td valign="top" align="left">Affection nostalgia</td>
<td valign="top" align="left">Readi, pretti, miss, yay, love readi pretti miss yay love song tomorrow goodnight saturday amp sleep hang alreadi night</td>
</tr>
<tr>
<td valign="top" align="left">7</td>
<td valign="top" align="left">Planning anticipation</td>
<td valign="top" align="left">Day, ll, someth, think, later day ll someth think later start make beauti tomorrow happen days amp let money</td>
</tr>
<tr>
<td valign="top" align="left">8</td>
<td valign="top" align="left">Work productivity</td>
<td valign="top" align="left">Work, glad, time, snow, home work glad time snow home hard today earli okay tonight easter night</td>
</tr>
<tr>
<td valign="top" align="left">9</td>
<td valign="top" align="left">Celebrations greetings</td>
<td valign="top" align="left">Happi, birthday, wonder, peopl, mani happi birthday wonder peopl mani love best sorri repli realli sooo</td>
</tr>
<tr>
<td valign="top" align="left">10</td>
<td valign="top" align="left">School life</td>
<td valign="top" align="left">Today, school, life, break, lunch today school life break lunch hour watch hear spring funni till wanna room</td>
</tr>
<tr>
<td valign="top" align="left">11</td>
<td valign="top" align="left">Quotations humor</td>
<td valign="top" align="left">Quot, listen, cold, like, hahaha quot listen cold like hahaha music babi love problem hey th great haha brother smile song</td>
</tr>
<tr>
<td valign="top" align="left">12</td>
<td valign="top" align="left">Love relationships</td>
<td valign="top" align="left">Love, tweet, nice, tire, amp love tweet nice tire amp awesome twitter summer train join kinda ddlovato</td>
</tr>
<tr>
<td valign="top" align="left">13</td>
<td valign="top" align="left">Happy</td>
<td valign="top" align="left">Awesom, Watch, like, send, wear awesom watch like send wear place love asot400 way help house man make amp fail</td>
</tr>
<tr>
<td valign="top" align="left">14</td>
<td valign="top" align="left">Technology</td>
<td valign="top" align="left">Know, need, want, twitter, hello know need want twitter hello dont think phone like love realli mileycyru fm cute</td>
</tr>
<tr>
<td valign="top" align="left">15</td>
<td valign="top" align="left">Reading learning</td>
<td valign="top" align="left">Read, alway, book, final, food read alway book final food time gone dinner believ think iphon famili tonight pick</td>
</tr>
<tr>
<td valign="top" align="left">16</td>
<td valign="top" align="left">Weekend activities</td>
<td valign="top" align="left">Good, weekend, girl, sick, luck good weekend girl sick luck night rain dream wish fuck shower tuesday great time today</td>
</tr>
<tr>
<td valign="top" align="left">17</td>
<td valign="top" align="left">Online engagement</td>
<td valign="top" align="left">Http, com, thank, follow, twitpic http com thank follow twitpic www tinyurl thanks check twitter link</td>
</tr>
<tr>
<td valign="top" align="left">18</td>
<td valign="top" align="left">Future plans</td>
<td valign="top" align="left">Time, watch, soon, movi, long time watch soon movi long suck come today love want realli heard movie anoth real</td>
</tr>
<tr>
<td valign="top" align="left">19</td>
<td valign="top" align="left">Sleep relaxation</td>
<td valign="top" align="left">Great, sleep, time, post, hope great sleep time post hope bore bit late everyth ly free breakfast http day</td>
</tr>
<tr>
<td valign="top" align="left">20</td>
<td valign="top" align="left">Visuals photos</td>
<td valign="top" align="left">Look, like, wait, forward, picture look like wait forward pictur sound realli think ll gt welcome twitter awww</td>
</tr></tbody>
</table>
</table-wrap>
<p>The results of research experiments are summarized in <xref ref-type="table" rid="T6">Table 6</xref>, which shows the accuracy and prediction probabilities for each combination of feature extraction method and classifier. Additionally, <xref ref-type="table" rid="T7">Table 7</xref> provides a comparative analysis of several ML classifiers, namely ANN, XGBoost, RF, and SVM, using different feature extraction techniques. These features include LDA, TF-IDF, N-gram, BOW, GloVe, and various combinations of these techniques. The classifiers are evaluated based on their performance metrics: Precision, Recall, and F1-score.</p>
<table-wrap position="float" id="T6">
<label>Table 6</label>
<caption><p>Results of classification performance ML models and feature selection methods based on accuracy.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>Sr</bold></th>
<th valign="top" align="left"><bold>Features</bold></th>
<th valign="top" align="center"><bold>ANN</bold></th>
<th valign="top" align="center"><bold>XGBoost</bold></th>
<th valign="top" align="center"><bold>RF</bold></th>
<th valign="top" align="center"><bold>SVM</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">1</td>
<td valign="top" align="left">LDA</td>
<td valign="top" align="center">72</td>
<td valign="top" align="center">68</td>
<td valign="top" align="center">72</td>
<td valign="top" align="center">72</td>
</tr>
<tr>
<td valign="top" align="left">2</td>
<td valign="top" align="left">TF-IDF</td>
<td valign="top" align="center">72</td>
<td valign="top" align="center">77</td>
<td valign="top" align="center">78</td>
<td valign="top" align="center">79</td>
</tr>
<tr>
<td valign="top" align="left">3</td>
<td valign="top" align="left">N-gram</td>
<td valign="top" align="center">73</td>
<td valign="top" align="center">78</td>
<td valign="top" align="center">76</td>
<td valign="top" align="center">77</td>
</tr>
<tr>
<td valign="top" align="left">4</td>
<td valign="top" align="left">BOW</td>
<td valign="top" align="center">73</td>
<td valign="top" align="center">78</td>
<td valign="top" align="center">75</td>
<td valign="top" align="center">78</td>
</tr>
<tr>
<td valign="top" align="left">5</td>
<td valign="top" align="left">GloVe</td>
<td valign="top" align="center">72</td>
<td valign="top" align="center">86</td>
<td valign="top" align="center">88</td>
<td valign="top" align="center">85</td>
</tr>
<tr>
<td valign="top" align="left">6</td>
<td valign="top" align="left">LDA&#x0002B;TFIDF&#x0002B;Ngram</td>
<td valign="top" align="center">78</td>
<td valign="top" align="center">77</td>
<td valign="top" align="center">72</td>
<td valign="top" align="center">72</td>
</tr>
<tr>
<td valign="top" align="left">7</td>
<td valign="top" align="left">LDA&#x0002B;BOW&#x0002B;TFIDF</td>
<td valign="top" align="center">76</td>
<td valign="top" align="center">77</td>
<td valign="top" align="center">77</td>
<td valign="top" align="center">78</td>
</tr></tbody>
</table>
</table-wrap>
<table-wrap position="float" id="T7">
<label>Table 7</label>
<caption><p>Comparison of ML models and feature selection methods based on additional metrics such as precision, recall, and F-score.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th/>
<th/>
<th valign="top" align="center" colspan="3"><bold>ANN</bold></th>
<th valign="top" align="center" colspan="3"><bold>XGBoost</bold></th>
<th valign="top" align="center" colspan="3"><bold>RF</bold></th>
<th valign="top" align="center" colspan="3"><bold>SVM</bold></th>
</tr>
<tr>
<th valign="top" align="left"><bold>Sr</bold></th>
<th valign="top" align="left"><bold>Features</bold></th>
<th valign="top" align="center"><bold>Prec</bold>.</th>
<th valign="top" align="center"><bold>Rec</bold>.</th>
<th valign="top" align="center"><bold>F1</bold></th>
<th valign="top" align="center"><bold>Prec</bold>.</th>
<th valign="top" align="center"><bold>Rec</bold>.</th>
<th valign="top" align="center"><bold>F1</bold></th>
<th valign="top" align="center"><bold>Prec</bold>.</th>
<th valign="top" align="center"><bold>Rec</bold>.</th>
<th valign="top" align="center"><bold>F1</bold></th>
<th valign="top" align="center"><bold>Prec</bold>.</th>
<th valign="top" align="center"><bold>Rec</bold>.</th>
<th valign="top" align="center"><bold>F1</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">1</td>
<td valign="top" align="left">LDA</td>
<td valign="top" align="center">73</td>
<td valign="top" align="center">94</td>
<td valign="top" align="center">82</td>
<td valign="top" align="center">74</td>
<td valign="top" align="center">87</td>
<td valign="top" align="center">80</td>
<td valign="top" align="center">74</td>
<td valign="top" align="center">94</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">72</td>
<td valign="top" align="center">98</td>
<td valign="top" align="center">83</td>
</tr>
<tr>
<td valign="top" align="left">2</td>
<td valign="top" align="left">TF-IDF</td>
<td valign="top" align="center">82</td>
<td valign="top" align="center">80</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">80</td>
<td valign="top" align="center">90</td>
<td valign="top" align="center">85</td>
<td valign="top" align="center">80</td>
<td valign="top" align="center">92</td>
<td valign="top" align="center">86</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">92</td>
<td valign="top" align="center">86</td>
</tr>
<tr>
<td valign="top" align="left">3</td>
<td valign="top" align="left">N-gram</td>
<td valign="top" align="center">83</td>
<td valign="top" align="center">79</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">91</td>
<td valign="top" align="center">86</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">86</td>
<td valign="top" align="center">84</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">94</td>
<td valign="top" align="center">86</td>
</tr>
<tr>
<td valign="top" align="left">4</td>
<td valign="top" align="left">BOW</td>
<td valign="top" align="center">83</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">85</td>
<td valign="top" align="center">81</td>
<td valign="top" align="center">91</td>
<td valign="top" align="center">85</td>
<td valign="top" align="center">82</td>
<td valign="top" align="center">84</td>
<td valign="top" align="center">83</td>
<td valign="top" align="center">79</td>
<td valign="top" align="center">94</td>
<td valign="top" align="center">86</td>
</tr>
<tr>
<td valign="top" align="left">5</td>
<td valign="top" align="left">GloVe</td>
<td valign="top" align="center">72</td>
<td valign="top" align="center">99</td>
<td valign="top" align="center">83</td>
<td valign="top" align="center">86</td>
<td valign="top" align="center">88</td>
<td valign="top" align="center">87</td>
<td valign="top" align="center">97</td>
<td valign="top" align="center">92</td>
<td valign="top" align="center">94</td>
<td valign="top" align="center">83</td>
<td valign="top" align="center">94</td>
<td valign="top" align="center">90</td>
</tr>
<tr>
<td valign="top" align="left">6</td>
<td valign="top" align="left">LDA&#x0002B;TFIDF&#x0002B;Ngram</td>
<td valign="top" align="center">80</td>
<td valign="top" align="center">91</td>
<td valign="top" align="center">85</td>
<td valign="top" align="center">82</td>
<td valign="top" align="center">87</td>
<td valign="top" align="center">84</td>
<td valign="top" align="center">79</td>
<td valign="top" align="center">92</td>
<td valign="top" align="center">84</td>
<td valign="top" align="center">79</td>
<td valign="top" align="center">92</td>
<td valign="top" align="center">84</td>
</tr>
<tr>
<td valign="top" align="left">7</td>
<td valign="top" align="left">LDA&#x0002B;BOW&#x0002B;TFIDF</td>
<td valign="top" align="center">82</td>
<td valign="top" align="center">85</td>
<td valign="top" align="center">84</td>
<td valign="top" align="center">78</td>
<td valign="top" align="center">94</td>
<td valign="top" align="center">86</td>
<td valign="top" align="center">77</td>
<td valign="top" align="center">95</td>
<td valign="top" align="center">85</td>
<td valign="top" align="center">79</td>
<td valign="top" align="center">94</td>
<td valign="top" align="center">86</td>
</tr></tbody>
</table>
</table-wrap>
<p>It should be noted from <xref ref-type="table" rid="T6">Table 6</xref> that the GloVe feature extraction technique combined with RF achieved the highest accuracy of 88%, demonstrating its ability to capture rich semantic information effectively. SVM also performed well with GloVe, achieving an accuracy of 85%. TF-IDF and N-gram modeling showed competitive performance, with XGB achieving an accuracy of 77% and 78%, respectively, demonstrating their effectiveness in capturing text features. BOW proved to be a reliable feature extraction method, with XGB and ANN yielding accuracies of 78% and 73%, respectively. The combination of LDA with TF-IDF and N-gram yielded an accuracy of 78% for ANN, showcasing the potential of combining multiple feature extraction techniques. Our experiments demonstrate the varied strengths of feature extraction methods and classifiers, with GloVe providing the most insightful semantic information. At the same time, N-gram and BOW maintained a consistent balance of performance across models.</p>
<p><xref ref-type="fig" rid="F4">Figure 4</xref> presents a comparative analysis of four ML classifiers ANNs, XGB, RF, and SVM across various feature extraction methods: LDA, TF-IDF, N-gram, BOW, GloVe, LDA&#x0002B;TF-IDF&#x0002B;Ngram, and LDA&#x0002B;BOW&#x0002B;TFIDF. The performance of each combination is evaluated and depicted in terms of accuracy percentages. The highest performance is achieved by the GloVe feature extraction method, followed closely by the SVM algorithm at 86%. Overall, the GloVe method consistently outperforms other feature extraction techniques, indicating its effectiveness in capturing word semantics and improving classification accuracy. Other notable performances include the LDA&#x0002B;TF-IDF&#x0002B;Ngram method, which demonstrates balanced accuracy across different classifiers, particularly with SVM classifiers. This comprehensive comparison underscores the importance of selecting suitable feature extraction methods to boost the predictive power of machine learning models in text classification tasks.</p>
<fig position="float" id="F4">
<label>Figure 4</label>
<caption><p>Comparison of ML models and feature selection methods based on accuracy.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-08-1627078-g0004.tif">
<alt-text>Bar chart comparing accuracy of different algorithms (ANN, XGB, RF, SVM) across text representations: LDA, TF-IDF, N-gram, BOW, GloVe, LDA&#x0002B;TF-IDF&#x0002B;N-gram, LDA&#x0002B;BOW&#x0002B;TF-IDF. Bars show higher accuracy for SVM and RF in most cases, especially for LDA&#x0002B;TF-IDF&#x0002B;N-gram and LDA&#x0002B;BOW&#x0002B;TF-IDF.</alt-text>
</graphic>
</fig>
<sec>
<title>5.1 Overview of evaluation</title>
<p>The experiments were performed on the same dataset for all feature extraction methods and classifiers to ensure consistent comparison. We evaluated seven feature extraction methods (LDA, TF-IDF, N-gram, BOW, GloVe, LDA&#x0002B;TFIDF&#x0002B;N-gram, and LDA&#x0002B;BOW&#x0002B;TFIDF) and four classifiers (ANN, XGBoost, RF, SVM). To improve analytical clarity, results are presented in two separate views: (1) ranking of feature extraction methods averaged across all classifiers, and (2) ranking of classifiers averaged across all feature extraction methods. Detailed per-configuration precision, recall and F1-score values are provided in the <xref ref-type="table" rid="T8">Tables 8</xref>, <xref ref-type="table" rid="T9">9</xref>.</p>
<table-wrap position="float" id="T8">
<label>Table 8</label>
<caption><p>Average accuracy of feature extraction methods across ML classifiers.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>Feature extraction</bold></th>
<th valign="top" align="center"><bold>Accuracy (%)</bold></th>
<th valign="top" align="center"><bold>Rank</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">GloVe</td>
<td valign="top" align="center">82.75</td>
<td valign="top" align="center">1</td>
</tr>
<tr>
<td valign="top" align="left">LDA &#x0002B; BOW &#x0002B; TF-IDF</td>
<td valign="top" align="center">77.00</td>
<td valign="top" align="center">2</td>
</tr>
<tr>
<td valign="top" align="left">TF-IDF</td>
<td valign="top" align="center">76.50</td>
<td valign="top" align="center">3</td>
</tr>
<tr>
<td valign="top" align="left">N-gram</td>
<td valign="top" align="center">76.00</td>
<td valign="top" align="center">4</td>
</tr>
<tr>
<td valign="top" align="left">BOW</td>
<td valign="top" align="center">76.00</td>
<td valign="top" align="center">5</td>
</tr>
<tr>
<td valign="top" align="left">LDA &#x0002B; TF-IDF &#x0002B; N-gram</td>
<td valign="top" align="center">74.75</td>
<td valign="top" align="center">6</td>
</tr>
<tr>
<td valign="top" align="left">LDA</td>
<td valign="top" align="center">71.00</td>
<td valign="top" align="center">7</td>
</tr></tbody>
</table>
</table-wrap>
<table-wrap position="float" id="T9">
<label>Table 9</label>
<caption><p>Average accuracy of ML classifiers across feature extraction methods.</p></caption>
<table frame="box" rules="all">
<thead>
<tr>
<th valign="top" align="left"><bold>Classifier</bold></th>
<th valign="top" align="center"><bold>Avg accuracy (%)</bold></th>
<th valign="top" align="center"><bold>Rank</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">XGBoost</td>
<td valign="top" align="center">77.29</td>
<td valign="top" align="center">1</td>
</tr>
<tr>
<td valign="top" align="left">SVM</td>
<td valign="top" align="center">77.29</td>
<td valign="top" align="center">2</td>
</tr>
<tr>
<td valign="top" align="left">Random Forest (RF)</td>
<td valign="top" align="center">76.86</td>
<td valign="top" align="center">3</td>
</tr>
<tr>
<td valign="top" align="left">ANN</td>
<td valign="top" align="center">73.71</td>
<td valign="top" align="center">4</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec>
<title>5.2 Feature extraction method rankings</title>
<p><xref ref-type="table" rid="T8">Table 8</xref> reports the average accuracy of each feature extraction method calculated over all four classifiers. This ranking shows which feature representations perform best on average in our study.</p>
<p>Interpretation: GloVe embeddings provide the highest average accuracy (82.75%) across classifiers, indicating that dense semantic representations capture context useful for depressive-linguistic signals in our dataset. Traditional lexical representations (TF-IDF, N-gram, BOW) remain competitive and may be preferable in resource-constrained settings.</p>
</sec>
<sec>
<title>5.3 Classifier rankings</title>
<p><xref ref-type="table" rid="T9">Table 9</xref> shows the average accuracy of each ML classifier across all feature extraction methods. This ranking isolates classifier performance independent of any single feature choice.</p>
</sec>
</sec>
<sec id="s6">
<title>6 LIME analysis</title>
<p>Despite their excellent performance, several ML models are often characterized as black boxes that produce outputs without offering explicit insights into the underlying reasoning behind their decisions. Understanding and interpreting the decision-making processes of such models is critical, particularly in applications where trust, transparency, and accountability are paramount. Consequently, it is imperative to examine the outputs of these models thoroughly and, more importantly, to develop methodologies that enable the generation of interpretable explanations for their decisions (<xref ref-type="bibr" rid="B91">S&#x000F8;gaard, 2021</xref>; <xref ref-type="bibr" rid="B34">Gohel et al., 2021</xref>). Providing explanations for a model&#x00027;s output enhances our ability to evaluate its predictions critically, thereby fostering greater confidence in determining whether to trust or question its outcomes.</p>
<p>To investigate the model&#x00027;s explainability, we employed LIME, a popular XAI technique that facilitates the interpretation of outputs without requiring direct inspection of the model&#x00027;s internal structure. LIME achieves this by perturbing the local features surrounding a specific target prediction and analyzing the corresponding changes in the model&#x00027;s output. In our experiments, the words surrounding a target entity were modified systematically, and the effects on the model&#x00027;s predictions were subsequently assessed to gain insights into the decision-making process.</p>
<p>Each subplot in <xref ref-type="fig" rid="F5">Figure 5</xref> illustrates the LIME analysis visualizations, providing an interpretable explanation of the predictions made by different classifiers for specific instances. Each subplot highlights the contribution of individual words (features) toward the prediction of either &#x0201C;Depression&#x0201D; or &#x0201C;Non-Depression&#x0201D; labels. The importance of the words is represented as bars, where positive contributions toward &#x0201C;Depression&#x0201D; are shown in blue, and contributions toward &#x0201C;Non-Depression&#x0201D; are shown in orange.</p>
<fig position="float" id="F5">
<label>Figure 5</label>
<caption><p>Example of explainable AI visualizations using LIME. <bold>(A)</bold> ANN and LDA feature extraction on 10th instance. <bold>(B)</bold> XGB and TF-IDF feature extraction on 960th instance. <bold>(C)</bold> SVM and BOW feature extraction on 1000th instance <bold>(D)</bold> RF and N gram feature extraction on 2000th instance. <bold>(E)</bold> ANN and GLOVE feature extraction on 3000th instance.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="frai-08-1627078-g0005.tif">
<alt-text>Five panels displaying analysis of text for depression versus non-depression using LIME across different models and feature extractions. Each panel shows a bar graph comparing scores for depression and non-depression, with highlighted influential words for each category. Models include ANN, LDA, XGB, TF-IDF, SVM, BOW, RF, and N Gram across various instances.</alt-text>
</graphic>
</fig>
<p><xref ref-type="fig" rid="F5">Figure 5a</xref> demonstrates LIME analysis for ANN with LDA features extraction on the 10th instance. It can be noted that the word &#x0201C;snow&#x0201D; contributed significantly to the &#x0201C;Depression&#x0201D; label, with a probability score of 0.54, compared to 0.46 for &#x0201C;Non-Depression.&#x0201D; This shows how specific context-sensitive words can influence predictions. Similarly, <xref ref-type="fig" rid="F5">Figure 5b</xref> demonstrates LIME analysis for XGB and TF-IDF on the 960th instance, the words &#x0201C;wish,&#x0201D; &#x0201C;sleep,&#x0201D; and &#x0201C;tooo&#x0201D; strongly contributed to the &#x0201C;Depression&#x0201D; label, achieving a high probability score of 0.87 for &#x0201C;Depression&#x0201D;. This highlights the model&#x00027;s ability to capture sentiment-relevant features using the TF-IDF approach.</p>
<p>In addition, <xref ref-type="fig" rid="F5">Figure 5c</xref> demonstrates a dominant contribution of the words &#x0201C;miss&#x0201D; and &#x0201C;sorry&#x0201D; toward the &#x0201C;Depression&#x0201D; label, achieving a probability score of 0.93. This emphasizes the BOW model&#x00027;s capability to detect sentiment-related words. Similarly, <xref ref-type="fig" rid="F5">Figure 5d</xref> visualizes that, RF with N-Gram on the 2000th instance, the words &#x0201C;work,&#x0201D; &#x0201C;desk,&#x0201D; and &#x0201C;10am&#x0201D; contribute strongly toward the &#x0201C;Depression&#x0201D; label, with a probability score of 0.91. This demonstrates how the N-Gram technique effectively captures contextual word co-occurrences. On the other hand, <xref ref-type="fig" rid="F5">Figure 5e</xref> indicates that ANN with GloVe embeddings on the 3000th instance, provided a contrasting prediction, with the word &#x0201C;islandiva147&#x0201D; contributing more toward &#x0201C;Non-Depression&#x0201D; than &#x0201C;Depression.&#x0201D; This suggests that GloVe embeddings capture semantic nuances but may misinterpret words out of context.</p>
<p>Overall, these visualizations illustrate the interpretability of the classifiers and their reliance on feature-specific contributions. While ANN and SVM demonstrated strong performance on LDA and BOW features, respectively, XGB and RF highlighted the importance of TF-IDF and N-Gram features. GloVe embeddings, while semantically rich, occasionally misinterpret specific instances, underscoring the importance of feature selection.</p>
</sec>
<sec sec-type="discussion" id="s7">
<title>7 Discussion</title>
<p>Our study examines various feature sets and classifiers, achieving notable accuracies, particularly when combining GloVe with classifiers such as RF and SVM. GloVe&#x0002B;RF, achieved 88% accuracy obtained by <xref ref-type="bibr" rid="B55">Ji et al. (2022a)</xref>. In comparison, <xref ref-type="bibr" rid="B8">Amanat et al. (2022a)</xref> reported accuracies up to 96.4% using a combination of TF-IDF, BOW, SVM, and RF. While our results for GloVe&#x0002B;RF (88%) and GloVe&#x0002B;SVM (85%) are slightly lower, they are still competitive given the different data sources and methodologies used. which is within the range reported by comparable studies on depression detection using traditional machine learning approaches. Variations in reported accuracy across the literature are largely attributable to differences in datasets, preprocessing steps, and feature-classifier combinations, making direct score comparisons less meaningful. Instead, the emphasis here is on showing that competitive performance can be achieved alongside enhanced interpretability.</p>
<p>Among feature extraction methods, GloVe embeddings achieved the highest average accuracy across classifiers (82.75%), followed by LDA&#x0002B;BOW&#x0002B;TF-IDF (77.00%) and TF-IDF (76.50%). This indicates that semantic embeddings like GloVe capture richer contextual information than purely lexical representations. These results are consistent with prior work&#x02014;for instance, <xref ref-type="bibr" rid="B99">Tong et al. (2022)</xref> reported 86% accuracy using GloVe-based features in a multi-classifier ensemble, whereas our GloVe&#x0002B;RF model reached 88% on this dataset. Although <xref ref-type="bibr" rid="B8">Amanat et al. (2022a)</xref> achieved higher performance (96.4%) with TF-IDF&#x0002B;BOW&#x0002B;SVM/RF, differences in datasets and preprocessing make direct comparisons indicative rather than conclusive. Notably, N-gram and BOW features, while ranking lower in average performance (76.00%), still matched or exceeded the accuracy of some deep learning models in the literature, such as the CNN (78%) and LSTM (80%) reported by <xref ref-type="bibr" rid="B28">de Souza et al. (2022)</xref>, demonstrating that simpler representations can be competitive for certain datasets.</p>
<p>Among ML classifiers, XGBoost and SVM obtained the highest average accuracy across feature sets (77.29%), closely followed by Random Forest (76.86%), with ANN ranking lowest (73.71%). This suggests that tree-based ensembles and margin-based classifiers are generally better suited to the depression detection task when trained on short, noisy social media text. These trends align with the findings of <xref ref-type="bibr" rid="B102">Wani et al. (2022)</xref>, where SVM achieved 71% and KNN 62% using N-gram features, and with <xref ref-type="bibr" rid="B55">Ji et al. (2022a)</xref>, who reported competitive performance with SVM on short-text datasets. In our experiments, ANN performed best with the combined LDA&#x0002B;TF-IDF&#x0002B;N-gram feature set (78%), which slightly exceeds the BiLSTM performance (79%) reported by <xref ref-type="bibr" rid="B41">Haque et al. (2022)</xref>, showing that under certain feature configurations, neural networks can still achieve strong results.</p>
<p>A key finding is that performance depends on the interaction between the feature set and the classifier. For example, while GloVe consistently ranks highest among features, its combination with RF (88%) and SVM (85%) outperformed its pairing with ANN (72%). Similarly, the LDA&#x0002B;TF-IDF&#x0002B;N-gram feature set worked particularly well with ANN (78%), but less so with RF (72%). These variations underscore the importance of evaluating both dimensions independently before selecting an optimal configuration.</p>
<p>Another distinctive aspect of this study is the systematic application of LIME across all feature-classifier combinations, rather than to a single model. LIME provided interpretable, instance-level explanations, identifying key linguistic cues such as first-person pronouns, negative emotion words, and self-referential phrases that heavily influenced depressive content predictions. This transparency is essential for building trust with mental health professionals and differentiates our work from most prior studies, where explainability is rarely addressed at this scale.</p>
<p>It should be noted that the accuracy of depression detection models heavily relies on the authenticity and consistency of users&#x00027; social media content. Factors such as bias for social desirability, stigma, and personal tendencies can influence the way users express themselves online, potentially leading to inaccuracies in detection (<xref ref-type="bibr" rid="B90">Shah et al., 2025</xref>). In general, detecting depression from social media posts inherently depends on the assumption that users share relevant depressive symptoms or emotional cues in their online interactions. However, not all individuals with depression disclose their condition or express depressive symptoms publicly on social media, which presents a notable limitation of computational models relying solely on social media text. This limitation leads to potential false negatives, where affected individuals who do not manifest depressive behavior online may be missed by such systems.</p>
<p>Similarly, another recent study (<xref ref-type="bibr" rid="B7">Aldkheel and Zhou, 2024</xref>) highlighted that social media detection methods that focus on observable linguistic, visual, and behavioral signals indicative of depression cannot account for users who do not publicly share or mask their symptoms due to privacy concerns, social stigma, or personal choice. Moreover, reliance on textual content alone restricts detection to expressed emotions and behaviors, which may not comprehensively represent every user&#x00027;s mental health status.</p>
<p>The necessity for combining social media analysis with clinical validation and offline data is emphasized to address these limitations and improve detection reliability. Verification of ground truth through clinical evaluations or integration of medical records along with social media data can help overcome false negatives caused by the absence of explicit online depressive expression (<xref ref-type="bibr" rid="B7">Aldkheel and Zhou, 2024</xref>). Thus, while social media-based depression detection offers valuable early screening potential, it cannot substitute for comprehensive clinical diagnosis and does not capture all cases, especially among users who do not disclose symptoms online.</p>
<p>GloVe word embeddings worked better with models like Random Forest and ANN because they capture the meaning and relationships between words. Unlike methods like TF-IDF or Bag of Words that just count word frequency, GloVe places similar words (like &#x0201C;sad&#x0201D; and &#x0201C;unhappy&#x0201D;) close together in a way that shows their meaning. This helps the models better tell the difference between depressive and non-depressive posts, especially since social media posts are usually short. Traditional models like SVM and Random Forest did well with TF-IDF because they can handle large sets of sparse features. However, the ANN model didn&#x00027;t perform as well with TF-IDF or BoW since those features don&#x00027;t carry the deeper meaning that neural networks are designed to learn from. When we used LIME to explain the model&#x00027;s decisions, we found that GloVe helped the models focus on important words like &#x0201C;lonely,&#x0201D; &#x0201C;worthless,&#x0201D; and &#x0201C;help&#x0201D;&#x02014; strong signs of depression. In contrast, TF-IDF often picked up common but less meaningful words, which made the learning less effective.</p>
<p>We used social media data from platform X (formerly Twitter) because it is public, real-time, and short in format making it ideal for spotting signs of mental health issues. Compared to sites like Reddit or medical records, Twitter has more variety in language, which helps our model work better in real-world situations (<xref ref-type="bibr" rid="B27">De Choudhury et al., 2013</xref>).</p>
<p>To get useful features from the text, we used both basic methods (like TF-IDF, LDA, N-gram, and Bag of Words) and word embeddings (like GloVe). While advanced models like BERT understand context better, they are slower and harder to explain. GloVe gave us good results without needing too much computing power, which is important if we want to use this in real-time systems (<xref ref-type="bibr" rid="B76">Pennington et al., 2014</xref>).</p>
<p>We picked simple models like SVM and Random Forest because they work well with smaller datasets, are faster to train, and are easier to understand&#x02014;especially when used with tools like LIME. These models also don&#x00027;t need powerful hardware and can be used in apps or cloud platforms. We used LIME to explain how our models make decisions. It works with many kinds of models and helps make the results clearer for doctors and other users. This is important for mental health tools, where we need to be careful and ethical in how results are used (<xref ref-type="bibr" rid="B80">Ribeiro et al., 2016b</xref>).</p>
<p>Overall, our findings demonstrate the effectiveness of combining traditional and advanced NLP techniques with robust classifiers to achieve competitive performance in text classification tasks. Our results indicate that, while advanced deep learning models are powerful, conventional methods and hybrid approaches can also achieve competitive accuracy, offering more interpretable and computationally efficient alternatives.</p>
<p>To address the identified research gaps in the literature, our study proposes an interpretable and comparative framework for depression detection using Twitter data, comprising four key components. First, multiple NLP feature representations are employed, including TF-IDF, BoW, LDA, N-grams, and GloVe embeddings, to effectively capture both statistical and semantic characteristics of textual content. Second, a range of ML classifiers, including SVM, RF, ANN, and XGBoost, were evaluated under identical experimental conditions to assess their predictive performance. Third, explainability was integrated into the framework using the LIME method, which highlights the contribution of individual features to classification outcomes, thereby enhancing transparency and trust in the model. Finally, a comparative evaluation is performed in which all models and feature extraction techniques are applied to the same dataset and evaluated using consistent metrics, accuracy, precision, recall, and F1 score, to ensure fairness, reproducibility, and applicability in the real world.</p>
</sec>
<sec sec-type="conclusions" id="s8">
<title>8 Conclusion</title>
<p>Mental illness is a prevalent social issue driven by socioeconomic, clinical, and individual risk factors, and the rise of social media has allowed the analysis of user-generated content for early detection of depression. In this work, we evaluated the effectiveness of various feature extraction methods and machine learning classifiers in detecting depression from X posts. Our findings demonstrate that social media data can be effectively utilized for mental health monitoring, with methods such as N-gram, BOW, and TF-IDF providing significant insights. In particular, the combination of TF-IDF with XGB achieved the highest precision of 87%, while the GloVe embeddings with RF reached an accuracy of 88% with lower interpretability. The use of LIME highlighted the importance of balancing accuracy and interpretability in model outcomes. Future research should focus on incorporating advanced deep learning models such as Transformers and BERT, developing real-time detection systems, integrating multimodal data, expanding analyses to additional social media platforms, and addressing ethical and privacy concerns in mental health monitoring.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s9">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/supplementary material, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec sec-type="author-contributions" id="s10">
<title>Author contributions</title>
<p>SH: Conceptualization, Formal analysis, Methodology, Validation, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. MN: Conceptualization, Methodology, Supervision, Validation, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. NA: Conceptualization, Supervision, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. MF: Funding acquisition, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. RN: Supervision, Validation, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing.</p>
</sec>
<sec sec-type="funding-information" id="s11">
<title>Funding</title>
<p>The author(s) declare that no financial support was received for the research and/or publication of this article.</p>
</sec>
<sec sec-type="COI-statement" id="conf1">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted without any commercial or financial relationships that could potentially create a conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="s12">
<title>Generative AI statement</title>
<p>The author(s) declare that no Gen AI was used in the creation of this manuscript.</p>
<p>Any alternative text (alt text) provided alongside figures in this article has been generated by Frontiers with the support of artificial intelligence and reasonable efforts have been made to ensure accuracy, including review by the authors wherever possible. If you identify any issues, please contact us.</p>
</sec>
<sec sec-type="disclaimer" id="s13">
<title>Publisher&#x00027;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<fn-group>
<fn id="fn0001"><p><sup>1</sup><ext-link ext-link-type="uri" xlink:href="https://www.kaggle.com/datasets/adedamolaajewole/training1600000processednoemoticon">https://www.kaggle.com/datasets/adedamolaajewole/training1600000processednoemoticon</ext-link></p></fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Abd Yusof</surname> <given-names>N. F.</given-names></name> <name><surname>Lin</surname> <given-names>C.</given-names></name> <name><surname>Guerin</surname> <given-names>F.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;Analysing the causes of depressed mood from depression vulnerable individuals,&#x0201D;</article-title> in <source>Proceedings of the International Workshop on Digital Disease Detection using Social Media 2017 (DDDSM-2017)</source> (<publisher-loc>Cham</publisher-loc>: <publisher-name>Springer</publisher-name>), <fpage>9</fpage>&#x02013;<lpage>17</lpage>.</citation>
</ref>
<ref id="B2">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>AbuRaed</surname> <given-names>A. G. T.</given-names></name> <name><surname>Prikryl</surname> <given-names>E. A.</given-names></name> <name><surname>Carenini</surname> <given-names>G.</given-names></name> <name><surname>Janjua</surname> <given-names>N. Z.</given-names></name></person-group> (<year>2024</year>). <article-title>Long COVID discourse in Canada, the United States, and Europe: topic modeling and sentiment analysis of twitter data</article-title>. <source>J. Med. Internet Res</source>. <volume>26</volume>:<fpage>e59425</fpage>. <pub-id pub-id-type="doi">10.2196/59425</pub-id><pub-id pub-id-type="pmid">39652387</pub-id></citation></ref>
<ref id="B3">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Adarsh</surname> <given-names>V.</given-names></name> <name><surname>Kumar</surname> <given-names>P. A.</given-names></name> <name><surname>Lavanya</surname> <given-names>V.</given-names></name> <name><surname>Gangadharan</surname> <given-names>G.</given-names></name></person-group> (<year>2023</year>). <article-title>Fair and explainable depression detection in social media</article-title>. <source>Inf. Process. Manag</source>. <volume>60</volume>:<fpage>103168</fpage>. <pub-id pub-id-type="doi">10.1016/j.ipm.2022.103168</pub-id></citation>
</ref>
<ref id="B4">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Akhtar</surname> <given-names>H. M. U.</given-names></name> <name><surname>Nauman</surname> <given-names>M.</given-names></name> <name><surname>Akhtar</surname> <given-names>N.</given-names></name> <name><surname>Hameed</surname> <given-names>M.</given-names></name> <name><surname>Hameed</surname> <given-names>S.</given-names></name> <name><surname>Tareen</surname> <given-names>M. Z.</given-names></name> <etal/></person-group>. (<year>2025</year>). <article-title>Mitigating cyber threats: machine learning and explainable AI for phishing detection</article-title>. <source>VFAST Trans. Softw. Eng</source>. <volume>13</volume>, <fpage>170</fpage>&#x02013;<lpage>195</lpage>. <pub-id pub-id-type="doi">10.21015/vtse.v13i2.2129</pub-id></citation>
</ref>
<ref id="B5">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Al Asad</surname> <given-names>N.</given-names></name> <name><surname>Pranto</surname> <given-names>M. A. M.</given-names></name> <name><surname>Islam</surname> <given-names>M. M.</given-names></name></person-group> (<year>2024</year>). <article-title>&#x0201C;Explainable deep learning for mental health detection from English and Arabic social media posts,&#x0201D;</article-title> in <source>ACM Transactions on Asian and Low-Resource Language Information Processing Volume 23</source> (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>ACM</publisher-name>).</citation>
</ref>
<ref id="B6">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Al Qudah</surname> <given-names>I.</given-names></name> <name><surname>Hashem</surname> <given-names>I.</given-names></name> <name><surname>Soufyane</surname> <given-names>A.</given-names></name> <name><surname>Chen</surname> <given-names>W.</given-names></name> <name><surname>Merabtene</surname> <given-names>T.</given-names></name></person-group> (<year>2022</year>). <article-title>&#x0201C;Applying latent Dirichlet allocation technique to classify topics on sustainability using Arabic text,&#x0201D;</article-title> in <source>Intelligent Computing: Proceedings of the 2022 Computing Conference, Volume 1</source> (<publisher-loc>Cham</publisher-loc>: <publisher-name>Springer</publisher-name>), <fpage>630</fpage>&#x02013;<lpage>638</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-031-10461-9_43</pub-id></citation>
</ref>
<ref id="B7">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Aldkheel</surname> <given-names>A.</given-names></name> <name><surname>Zhou</surname> <given-names>L.</given-names></name></person-group> (<year>2024</year>). <article-title>Depression detection on social media: a classification framework and research challenges and opportunities. <italic>J. Healthc. Inform</italic></article-title>. <source>Res</source>. <volume>8</volume>, <fpage>88</fpage>&#x02013;<lpage>120</lpage>. <pub-id pub-id-type="doi">10.1007/s41666-023-00152-3</pub-id><pub-id pub-id-type="pmid">38273983</pub-id></citation></ref>
<ref id="B8">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Amanat</surname> <given-names>A.</given-names></name> <name><surname>Rizwan</surname> <given-names>M.</given-names></name> <name><surname>Javed</surname> <given-names>A. R.</given-names></name> <name><surname>Abdelhaq</surname> <given-names>M.</given-names></name> <name><surname>Alsaqour</surname> <given-names>R.</given-names></name> <name><surname>Pandya</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2022a</year>). <article-title>Deep learning for depression detection from textual data</article-title>. <source>Electronics</source> <volume>11</volume>:<fpage>676</fpage>. <pub-id pub-id-type="doi">10.3390/electronics11050676</pub-id></citation>
</ref>
<ref id="B9">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Amanat</surname> <given-names>A.</given-names></name> <name><surname>Shah</surname> <given-names>S. A. A.</given-names></name> <name><surname>Javaid</surname> <given-names>N.</given-names></name> <name><surname>Iqbal</surname> <given-names>F.</given-names></name></person-group> (<year>2022b</year>). <article-title>Detection of depression using convolutional neural networks and word2vec on reddit posts</article-title>. <source>IEEE Access</source> <volume>10</volume>, <fpage>14638</fpage>&#x02013;<lpage>14648</lpage>. <pub-id pub-id-type="doi">10.1080/23311975.2022.2039087</pub-id></citation>
</ref>
<ref id="B10">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ancona</surname> <given-names>M.</given-names></name> <name><surname>Ceolini</surname> <given-names>E.</given-names></name> <name><surname>&#x000D6;ztireli</surname> <given-names>C.</given-names></name> <name><surname>Gross</surname> <given-names>M.</given-names></name></person-group> (<year>2018</year>). <article-title>Towards better understanding of gradient-based attribution methods for deep neural networks</article-title>. <source>arXiv [Preprint]</source> arXiv:1711.06104. <pub-id pub-id-type="doi">10.48550/arXiv.1711.06104</pub-id></citation>
</ref>
<ref id="B11">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bojanowski</surname> <given-names>P.</given-names></name> <name><surname>Grave</surname> <given-names>E.</given-names></name> <name><surname>Joulin</surname> <given-names>A.</given-names></name> <name><surname>Mikolov</surname> <given-names>T.</given-names></name></person-group> (<year>2017</year>). <article-title>Enriching word vectors with subword information</article-title>. <source>Trans. Assoc. Comput. Linguist</source>. <volume>5</volume>, <fpage>135</fpage>&#x02013;<lpage>146</lpage>. <pub-id pub-id-type="doi">10.1162/tacl_a_00051</pub-id></citation>
</ref>
<ref id="B12">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Breiman</surname> <given-names>L.</given-names></name></person-group> (<year>2001</year>). <article-title>Random forests</article-title>. <source>Mach. Learn</source>. <volume>45</volume>, <fpage>5</fpage>&#x02013;<lpage>32</lpage>. <pub-id pub-id-type="doi">10.1023/A:1010933404324</pub-id></citation>
</ref>
<ref id="B13">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Brochier</surname> <given-names>R.</given-names></name> <name><surname>Guille</surname> <given-names>A.</given-names></name> <name><surname>Velcin</surname> <given-names>J.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Global vectors for node representations,&#x0201D;</article-title> in <source>The World Wide Web Conference</source> (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>ACM</publisher-name>), <fpage>2587</fpage>&#x02013;<lpage>2593</lpage>. <pub-id pub-id-type="doi">10.1145/3308558.3313595</pub-id></citation>
</ref>
<ref id="B14">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cambria</surname> <given-names>E.</given-names></name> <name><surname>White</surname> <given-names>B.</given-names></name></person-group> (<year>2014</year>). <article-title>Jumping NLP curves: a review of natural language processing research</article-title>. <source>IEEE Comput. Intell. Mag</source>. <volume>9</volume>, <fpage>48</fpage>&#x02013;<lpage>57</lpage>. <pub-id pub-id-type="doi">10.1109/MCI.2014.2307227</pub-id></citation>
</ref>
<ref id="B15">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cesarini</surname> <given-names>M.</given-names></name> <name><surname>Malandri</surname> <given-names>L.</given-names></name> <name><surname>Pallucchini</surname> <given-names>F.</given-names></name> <name><surname>Seveso</surname> <given-names>A.</given-names></name> <name><surname>Xing</surname> <given-names>F.</given-names></name></person-group> (<year>2024</year>). <article-title>Explainable AI for text classification: lessons from a comprehensive evaluation of <italic>post hoc</italic> methods</article-title>. <source>Cogn. Comput</source>. <volume>16</volume>, <fpage>3077</fpage>&#x02013;<lpage>3095</lpage>. <pub-id pub-id-type="doi">10.1007/s12559-024-10325-w</pub-id></citation>
</ref>
<ref id="B16">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cha</surname> <given-names>J.</given-names></name> <name><surname>Kim</surname> <given-names>S.</given-names></name> <name><surname>Kim</surname> <given-names>D.</given-names></name> <name><surname>Park</surname> <given-names>E.</given-names></name></person-group> (<year>2024</year>). <article-title>Mogam: a multimodal object-oriented graph attention model for depression detection</article-title>. <source>arXiv [Preprint]</source>. arXiv:2403.15485. <pub-id pub-id-type="doi">10.48550/arXiv.2403.15485</pub-id></citation>
</ref>
<ref id="B17">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chakraborty</surname> <given-names>D.</given-names></name> <name><surname>Ivan</surname> <given-names>C.</given-names></name> <name><surname>Amero</surname> <given-names>P.</given-names></name> <name><surname>Khan</surname> <given-names>M.</given-names></name> <name><surname>Rodriguez-Aguayo</surname> <given-names>C.</given-names></name> <name><surname>Ba&#x0015F;a&#x0011F;ao&#x0011F;lu</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Explainable artificial intelligence reveals novel insight into tumor microenvironment conditions linked with better prognosis in patients with breast cancer</article-title>. <source>Cancers</source> <volume>13</volume>:<fpage>3450</fpage>. <pub-id pub-id-type="doi">10.3390/cancers13143450</pub-id><pub-id pub-id-type="pmid">34298668</pub-id></citation></ref>
<ref id="B18">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chancellor</surname> <given-names>S.</given-names></name> <name><surname>De Choudhury</surname> <given-names>M.</given-names></name></person-group> (<year>2020</year>). <article-title>Methods in predictive techniques for mental health status on social media: a critical review</article-title>. <source>NPJ Digit. Med</source>. <volume>3</volume>, <fpage>1</fpage>&#x02013;<lpage>11</lpage>. <pub-id pub-id-type="doi">10.1038/s41746-020-0233-7</pub-id><pub-id pub-id-type="pmid">32219184</pub-id></citation></ref>
<ref id="B19">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chanda</surname> <given-names>K.</given-names></name> <name><surname>Roy</surname> <given-names>S.</given-names></name> <name><surname>Mondal</surname> <given-names>H.</given-names></name> <name><surname>Bose</surname> <given-names>R.</given-names></name></person-group> (<year>2022</year>). <article-title>To judge depression and mental illness on social media using twitter</article-title>. <source>Univers. J. Public Health</source> <volume>10</volume>, <fpage>116</fpage>&#x02013;<lpage>129</lpage>. <pub-id pub-id-type="doi">10.13189/ujph.2022.100113</pub-id><pub-id pub-id-type="pmid">30947635</pub-id></citation></ref>
<ref id="B20">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chelgani</surname> <given-names>S. C.</given-names></name> <name><surname>Nasiri</surname> <given-names>H.</given-names></name> <name><surname>Tohry</surname> <given-names>A.</given-names></name> <name><surname>Heidari</surname> <given-names>H.</given-names></name></person-group> (<year>2023</year>). <article-title>Modeling industrial hydrocyclone operational variables by shap-catboost-a &#x0201C;conscious lab&#x0201D; approach</article-title>. <source>Powder Technol</source>. <volume>420</volume>:<fpage>118416</fpage>. <pub-id pub-id-type="doi">10.1016/j.powtec.2023.118416</pub-id></citation>
</ref>
<ref id="B21">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>D.</given-names></name> <name><surname>Zhao</surname> <given-names>H.</given-names></name> <name><surname>He</surname> <given-names>J.</given-names></name> <name><surname>Pan</surname> <given-names>Q.</given-names></name> <name><surname>Zhao</surname> <given-names>W.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;An causal XAI diagnostic model for breast cancer based on mammography reports,&#x0201D;</article-title> in <source>2021 IEEE International Conference on Bioinformatics and Biomedicine (BIBM)</source> (<publisher-loc>Houston, TX</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>3341</fpage>&#x02013;<lpage>3349</lpage>. <pub-id pub-id-type="doi">10.1109/BIBM52615.2021.9669648</pub-id></citation>
</ref>
<ref id="B22">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>T.</given-names></name> <name><surname>Guestrin</surname> <given-names>C.</given-names></name></person-group> (<year>2016</year>). <article-title>&#x0201C;Xgboost: a scalable tree boosting system,&#x0201D;</article-title> in <source>Proceedings of the 22nd ACM SIGKDD International Conference on Knowledge Discovery and Data Mining</source> (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>ACM</publisher-name>), <fpage>785</fpage>&#x02013;<lpage>794</lpage>. <pub-id pub-id-type="doi">10.1145/2939672.2939785</pub-id></citation>
</ref>
<ref id="B23">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>T.</given-names></name> <name><surname>He</surname> <given-names>T.</given-names></name> <name><surname>Benesty</surname> <given-names>M.</given-names></name> <name><surname>Khotilovich</surname> <given-names>V.</given-names></name> <name><surname>Tang</surname> <given-names>Y.</given-names></name> <name><surname>Cho</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2015</year>). <article-title>&#x0201C;Xgboost: a scalable tree boosting system,&#x0201D;</article-title> in <source>Proceedings of the 25th International Conference on Big Data</source> (<publisher-loc>Cham</publisher-loc>: <publisher-name>Springer</publisher-name>).</citation>
</ref>
<ref id="B24">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>X.</given-names></name> <name><surname>Lin</surname> <given-names>X.</given-names></name></person-group> (<year>2025</year>). <article-title>Generating medically-informed explanations for depression detection using LLMS</article-title>. <source>arXiv [Preprint]</source>. arXiv:2503.14671. <pub-id pub-id-type="doi">10.48500/arXiv.2503.14671</pub-id></citation>
</ref>
<ref id="B25">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cho</surname> <given-names>G.</given-names></name> <name><surname>Yim</surname> <given-names>J.</given-names></name> <name><surname>Choi</surname> <given-names>Y.</given-names></name> <name><surname>Ko</surname> <given-names>J.</given-names></name> <name><surname>Lee</surname> <given-names>S.-H.</given-names></name></person-group> (<year>2019</year>). <article-title>Review of machine learning algorithms for diagnosing mental illness</article-title>. <source>Psychiatry Investig</source>. <volume>16</volume>:<fpage>262</fpage>. <pub-id pub-id-type="doi">10.30773/pi.2018.12.21.2</pub-id><pub-id pub-id-type="pmid">30947496</pub-id></citation></ref>
<ref id="B26">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cortes</surname> <given-names>C.</given-names></name> <name><surname>Vapnik</surname> <given-names>V.</given-names></name></person-group> (<year>1995</year>). <article-title>Support-vector networks</article-title>. <source>Mach. Learn</source>. <volume>20</volume>, <fpage>273</fpage>&#x02013;<lpage>297</lpage>. <pub-id pub-id-type="doi">10.1023/A:1022627411411</pub-id></citation>
</ref>
<ref id="B27">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>De Choudhury</surname> <given-names>M.</given-names></name> <name><surname>Gamon</surname> <given-names>M.</given-names></name> <name><surname>Counts</surname> <given-names>S.</given-names></name> <name><surname>Horvitz</surname> <given-names>E.</given-names></name></person-group> (<year>2013</year>). <article-title>&#x0201C;Predicting depression via social media,&#x0201D;</article-title> in <source>Proceedings of the Seventh International AAAI Conference on Weblogs and Social Media</source>.</citation>
</ref>
<ref id="B28">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>de Souza</surname> <given-names>V. B.</given-names></name> <name><surname>Nobre</surname> <given-names>J. C.</given-names></name> <name><surname>Becker</surname> <given-names>K.</given-names></name></person-group> (<year>2022</year>). <article-title>DAC stacking: a deep learning ensemble to classify anxiety, depression, and their comorbidity from reddit texts</article-title>. <source>IEEE J. Biomed. Health Inform</source>. <volume>26</volume>, <fpage>3303</fpage>&#x02013;<lpage>3311</lpage>. <pub-id pub-id-type="doi">10.1109/JBHI.2022.3151589</pub-id><pub-id pub-id-type="pmid">35230959</pub-id></citation></ref>
<ref id="B29">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Devlin</surname> <given-names>J.</given-names></name> <name><surname>Chang</surname> <given-names>M.-W.</given-names></name> <name><surname>Lee</surname> <given-names>K.</given-names></name> <name><surname>Toutanova</surname> <given-names>K.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;Bert: pre-training of deep bidirectional transformers for language understanding,&#x0201D;</article-title> in <source>Proceedings of the 2019 conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies, Volume 1 (Long and Short Papers)</source> (<publisher-loc>Minneapolis, MN</publisher-loc>), <fpage>4171</fpage>&#x02013;<lpage>4186</lpage>.</citation>
</ref>
<ref id="B30">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dey</surname> <given-names>R. K.</given-names></name> <name><surname>Das</surname> <given-names>A. K.</given-names></name></person-group> (<year>2023</year>). <article-title>Modified term frequency-inverse document frequency based deep hybrid framework for sentiment analysis</article-title>. <source>Multimed. Tools Appl</source>. <volume>12</volume>, <fpage>1</fpage>&#x02013;<lpage>24</lpage>. <pub-id pub-id-type="doi">10.1007/s11042-023-14653-1</pub-id><pub-id pub-id-type="pmid">37362742</pub-id></citation></ref>
<ref id="B31">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fatima</surname> <given-names>B.</given-names></name> <name><surname>Amina</surname> <given-names>M.</given-names></name> <name><surname>Nachida</surname> <given-names>R.</given-names></name> <name><surname>Hamza</surname> <given-names>H.</given-names></name></person-group> (<year>2020</year>). <article-title>A mixed deep learning based model to early detection of depression</article-title>. <source>J. Web Eng</source>. <volume>19</volume>, <fpage>429</fpage>&#x02013;<lpage>455</lpage>. <pub-id pub-id-type="doi">10.13052/jwe1540-9589.19344</pub-id></citation>
</ref>
<ref id="B32">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Friedman</surname> <given-names>S.</given-names></name> <name><surname>Allin</surname> <given-names>L.</given-names></name> <name><surname>Gray</surname> <given-names>W.</given-names></name></person-group> (<year>2024</year>). <article-title>&#x0201C;The world is your oyster&#x0201D;: mothers&#x00027; perspectives on the value and purpose of an independent Forest School provision</article-title>. <source>Child. Geogr.</source> <volume>22</volume>, <fpage>597</fpage>&#x02013;<lpage>611</lpage>. <pub-id pub-id-type="doi">10.1080/14733285.2024.2321387</pub-id></citation>
</ref>
<ref id="B33">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ghosh</surname> <given-names>S.</given-names></name> <name><surname>Anwar</surname> <given-names>T.</given-names></name></person-group> (<year>2021</year>). <article-title>Depression intensity estimation via social media: a deep learning approach</article-title>. <source>IEEE Trans. Comput. Soc. Syst</source>. <volume>8</volume>, <fpage>1465</fpage>&#x02013;<lpage>1474</lpage>. <pub-id pub-id-type="doi">10.1109/TCSS.2021.3084154</pub-id></citation>
</ref>
<ref id="B34">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gohel</surname> <given-names>P.</given-names></name> <name><surname>Singh</surname> <given-names>P.</given-names></name> <name><surname>Mohanty</surname> <given-names>M.</given-names></name></person-group> (<year>2021</year>). <article-title>Explainable AI: current status and future directions</article-title>. <source>arXiv [Preprint]</source>. arXiv:2107.07045. <pub-id pub-id-type="doi">10.48550/arXiv.2107.07045</pub-id></citation>
</ref>
<ref id="B35">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Goodfellow</surname> <given-names>I.</given-names></name> <name><surname>Bengio</surname> <given-names>Y.</given-names></name> <name><surname>Courville</surname> <given-names>A.</given-names></name></person-group> (<year>2016</year>). <source>Deep Learning</source>. <publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>MIT Press</publisher-name>.</citation>
</ref>
<ref id="B36">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Govindasamy</surname> <given-names>K. A.</given-names></name> <name><surname>Palanichamy</surname> <given-names>N.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;Depression detection using machine learning techniques on twitter data,&#x0201D;</article-title> in <source>2021 5th International Conference on Intelligent Computing and Control Systems (ICICCS)</source> (<publisher-loc>Madurai</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>960</fpage>&#x02013;<lpage>966</lpage>. <pub-id pub-id-type="doi">10.1109/ICICCS51141.2021.9432203</pub-id></citation>
</ref>
<ref id="B37">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Guntuku</surname> <given-names>S. C.</given-names></name> <name><surname>Yaden</surname> <given-names>D. B.</given-names></name> <name><surname>Kern</surname> <given-names>M. L.</given-names></name> <name><surname>Ungar</surname> <given-names>L. H.</given-names></name> <name><surname>Eichstaedt</surname> <given-names>J. C.</given-names></name></person-group> (<year>2017</year>). <article-title>Detecting depression and mental illness on social media: an integrative review</article-title>. <source>Curr. Opin. Behav. Sci</source>. <volume>18</volume>, <fpage>43</fpage>&#x02013;<lpage>49</lpage>. <pub-id pub-id-type="doi">10.1016/j.cobeha.2017.07.005</pub-id></citation>
</ref>
<ref id="B38">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>M.</given-names></name> <name><surname>Zhang</surname> <given-names>H.</given-names></name></person-group> (<year>2023a</year>). <article-title>Lime-based explainability for mental health predictions on social media</article-title>. <source>J. Med. Internet Res</source>. <volume>25</volume>:<fpage>e43915</fpage>. <pub-id pub-id-type="doi">10.1016/j.scitotenv.2023.166118</pub-id><pub-id pub-id-type="pmid">40704178</pub-id></citation></ref>
<ref id="B39">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>Z.</given-names></name> <name><surname>Xu</surname> <given-names>X.</given-names></name></person-group> (<year>2023b</year>). <article-title>Research on the detection model of mental illness of online forum users based on convolutional network</article-title>. <source>BMC Psychol</source>. <volume>11</volume>:<fpage>424</fpage>. <pub-id pub-id-type="doi">10.1186/s40359-023-01460-4</pub-id><pub-id pub-id-type="pmid">38049891</pub-id></citation></ref>
<ref id="B40">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Hakkoum</surname> <given-names>H.</given-names></name> <name><surname>Idri</surname> <given-names>A.</given-names></name> <name><surname>Abnane</surname> <given-names>I.</given-names></name></person-group> (<year>2020</year>). <article-title>&#x0201C;Artificial neural networks interpretation using lime for breast cancer diagnosis,&#x0201D;</article-title> in <source>Trends and Innovations in Information Systems and Technologies: Volume 38</source> (<publisher-loc>Cham</publisher-loc>: <publisher-name>Springer</publisher-name>), <fpage>15</fpage>&#x02013;<lpage>24</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-030-45697-9_2</pub-id></citation>
</ref>
<ref id="B41">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Haque</surname> <given-names>R.</given-names></name> <name><surname>Islam</surname> <given-names>N.</given-names></name> <name><surname>Islam</surname> <given-names>M.</given-names></name> <name><surname>Ahsan</surname> <given-names>M. M.</given-names></name></person-group> (<year>2022</year>). <article-title>A comparative analysis on suicidal ideation detection using nlp, machine, and deep learning</article-title>. <source>Technologies</source> <volume>10</volume>:<fpage>57</fpage>. <pub-id pub-id-type="doi">10.3390/technologies10030057</pub-id></citation>
</ref>
<ref id="B42">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Harris</surname> <given-names>Z. S.</given-names></name></person-group> (<year>1954</year>). <article-title>Distributional structure</article-title>. <source>Word</source> <volume>10</volume>, <fpage>146</fpage>&#x02013;<lpage>162</lpage>. <pub-id pub-id-type="doi">10.1080/00437956.1954.11659520</pub-id></citation>
</ref>
<ref id="B43">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Helmy</surname> <given-names>A.</given-names></name> <name><surname>Nassar</surname> <given-names>R.</given-names></name> <name><surname>Ramdan</surname> <given-names>N.</given-names></name></person-group> (<year>2024</year>). <article-title>Depression detection for twitter users using sentiment analysis in English and Arabic tweets</article-title>. <source>Artif. Intell. Med</source>. <volume>147</volume>:<fpage>102716</fpage>. <pub-id pub-id-type="doi">10.1016/j.artmed.2023.102716</pub-id><pub-id pub-id-type="pmid">38184345</pub-id></citation></ref>
<ref id="B44">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Hemmatirad</surname> <given-names>K.</given-names></name> <name><surname>Bagherzadeh</surname> <given-names>H.</given-names></name> <name><surname>Fazl-Ersi</surname> <given-names>E.</given-names></name> <name><surname>Vahedian</surname> <given-names>A.</given-names></name></person-group> (<year>2020</year>). <article-title>&#x0201C;Detection of mental illness risk on social media through multi-level SVMS,&#x0201D;</article-title> in <source>2020 8th Iranian Joint Congress on Fuzzy and Intelligent Systems (CFIS)</source> (<publisher-loc>Mashhad</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>116</fpage>&#x02013;<lpage>120</lpage>. <pub-id pub-id-type="doi">10.1109/CFIS49607.2020.9238692</pub-id></citation>
</ref>
<ref id="B45">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ho</surname> <given-names>T. K.</given-names></name></person-group> (<year>1998</year>). <article-title>The random subspace method for constructing decision forests</article-title>. <source>IEEE Trans. Pattern Anal. Mach. Intell</source>. <volume>20</volume>, <fpage>832</fpage>&#x02013;<lpage>844</lpage>. <pub-id pub-id-type="doi">10.1109/34.709601</pub-id></citation>
</ref>
<ref id="B46">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hoque</surname> <given-names>A. M.</given-names></name> <name><surname>Rahman</surname> <given-names>A.</given-names></name> <name><surname>Hossain</surname> <given-names>M. E.</given-names></name> <name><surname>Hossain</surname> <given-names>M. S.</given-names></name> <name><surname>Muhammad</surname> <given-names>G.</given-names></name></person-group> (<year>2025</year>). <article-title>Effective depression detection and interpretation: integrating machine learning, deep learning, language models, and explainable AI</article-title>. <source>ARRAY</source> <volume>25</volume>:<fpage>100375</fpage>. <pub-id pub-id-type="doi">10.1016/j.array.2025.100375</pub-id></citation>
</ref>
<ref id="B47">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ibrahimov</surname> <given-names>K.</given-names></name> <name><surname>Ali</surname> <given-names>S.</given-names></name></person-group> (<year>2024</year>). <article-title>Importance of explainable artificial intelligence in mental health applications</article-title>. <source>Front. Artif. Intell</source>. <volume>7</volume>, <fpage>1</fpage>&#x02013;<lpage>12</lpage>.</citation>
</ref>
<ref id="B48">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ibrahimov</surname> <given-names>Y.</given-names></name> <name><surname>Anwar</surname> <given-names>T.</given-names></name> <name><surname>Yuan</surname> <given-names>T.</given-names></name></person-group> (<year>2024</year>). <article-title>Explainable AI for mental disorder detection via social media: a survey and outlook</article-title>. <source>arXiv [Preprint]</source>. arXiv:2406.05984. <pub-id pub-id-type="doi">10.4820/arXiv.2406.05984</pub-id></citation>
</ref>
<ref id="B49">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ibrahimov</surname> <given-names>Y.</given-names></name> <name><surname>Anwar</surname> <given-names>T.</given-names></name> <name><surname>Yuan</surname> <given-names>T.</given-names></name></person-group> (<year>2025</year>). <article-title>Depressionx: knowledge infused residual attention for explainable depression severity assessment</article-title>. <source>arXiv [Preprint]</source>. arXiv:2501.14985. <pub-id pub-id-type="doi">10.48550/arXiv.2501.14985</pub-id></citation>
</ref>
<ref id="B50">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Islam</surname> <given-names>M. R.</given-names></name> <name><surname>Kabir</surname> <given-names>M. A.</given-names></name> <name><surname>Ahmed</surname> <given-names>A.</given-names></name> <name><surname>Kamal</surname> <given-names>A. R. M.</given-names></name> <name><surname>Wang</surname> <given-names>H.</given-names></name> <name><surname>Ulhaq</surname> <given-names>A.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Depression detection from social network data using machine learning techniques</article-title>. <source>Health Inform. Sci. Syst</source>. <volume>6</volume>, <fpage>1</fpage>&#x02013;<lpage>12</lpage>. <pub-id pub-id-type="doi">10.1007/s13755-018-0046-0</pub-id><pub-id pub-id-type="pmid">30186594</pub-id></citation></ref>
<ref id="B51">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ive</surname> <given-names>J.</given-names></name> <name><surname>Viani</surname> <given-names>N.</given-names></name> <name><surname>Kam</surname> <given-names>J.</given-names></name> <name><surname>Yin</surname> <given-names>L.</given-names></name> <name><surname>Verma</surname> <given-names>S.</given-names></name> <name><surname>Puntis</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>Generation and evaluation of artificial mental health records for natural language processing</article-title>. <source>NPJ Digit. Med</source>. <volume>3</volume>, <fpage>1</fpage>&#x02013;<lpage>9</lpage>. <pub-id pub-id-type="doi">10.1038/s41746-020-0267-x</pub-id><pub-id pub-id-type="pmid">32435697</pub-id></citation></ref>
<ref id="B52">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jahromi</surname> <given-names>M. N.</given-names></name> <name><surname>Muddamsetty</surname> <given-names>S. M.</given-names></name> <name><surname>Jarlner</surname> <given-names>A. S. S.</given-names></name> <name><surname>H&#x000F8;genhaug</surname> <given-names>A. M.</given-names></name> <name><surname>Gammeltoft-Hansen</surname> <given-names>T.</given-names></name> <name><surname>Moeslund</surname> <given-names>T. B.</given-names></name></person-group> (<year>2024</year>). <article-title>Sidu-txt: an xai algorithm for nlp with a holistic assessment approach</article-title>. <source>Nat. Lang. Process. J</source>. <volume>7</volume>:<fpage>100078</fpage>. <pub-id pub-id-type="doi">10.1016/j.nlp.2024.100078</pub-id></citation>
</ref>
<ref id="B53">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jain</surname> <given-names>A.</given-names></name> <name><surname>Patel</surname> <given-names>H.</given-names></name> <name><surname>Nagalapatti</surname> <given-names>L.</given-names></name> <name><surname>Gupta</surname> <given-names>N.</given-names></name> <name><surname>Mehta</surname> <given-names>S.</given-names></name> <name><surname>Guttula</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>&#x0201C;Overview and importance of data quality for machine learning tasks,&#x0201D;</article-title> in <source>Proceedings of the 26th ACM SIGKDD International Conference on Knowledge Discovery &#x00026; Data Mining</source> (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>ACM</publisher-name>), <fpage>3561</fpage>&#x02013;<lpage>3562</lpage>. <pub-id pub-id-type="doi">10.1145/3394486.3406477</pub-id></citation>
</ref>
<ref id="B54">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jain</surname> <given-names>P.</given-names></name> <name><surname>Srinivas</surname> <given-names>K. R.</given-names></name> <name><surname>Vichare</surname> <given-names>A.</given-names></name></person-group> (<year>2022</year>). <article-title>Depression and suicide analysis using machine learning and nlp</article-title>. <source>J. Phys. Conf. Seri</source>. <volume>2161</volume>:<fpage>012034</fpage>. <pub-id pub-id-type="doi">10.1088/1742-6596/2161/1/012034</pub-id></citation>
</ref>
<ref id="B55">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ji</surname> <given-names>S.</given-names></name> <name><surname>Li</surname> <given-names>X.</given-names></name> <name><surname>Huang</surname> <given-names>Z.</given-names></name> <name><surname>Cambria</surname> <given-names>E.</given-names></name></person-group> (<year>2022a</year>). <article-title>Suicidal ideation and mental disorder detection with attentive relation networks</article-title>. <source>Neural Comput. Appl</source>. <volume>34</volume>, <fpage>10309</fpage>&#x02013;<lpage>10319</lpage>. <pub-id pub-id-type="doi">10.1007/s00521-021-06208-y</pub-id></citation>
</ref>
<ref id="B56">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ji</surname> <given-names>S.</given-names></name> <name><surname>Yu</surname> <given-names>C.</given-names></name> <name><surname>Fung</surname> <given-names>S. F.</given-names></name></person-group> (<year>2022b</year>). <article-title>Detecting depression on social media using Bert-based models: a case study on Twitter</article-title>. <source>J. Affect. Disord. Rep</source>. <volume>8</volume>:<fpage>100313</fpage>.</citation>
</ref>
<ref id="B57">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Joshi</surname> <given-names>B.</given-names></name> <name><surname>Shah</surname> <given-names>N.</given-names></name> <name><surname>Barbieri</surname> <given-names>F.</given-names></name> <name><surname>Neves</surname> <given-names>L.</given-names></name></person-group> (<year>2020</year>). <article-title>The devil is in the details: evaluating limitations of transformer-based methods for granular tasks</article-title>. <source>arXiv [Preprint]</source>. arXiv:2011.01196. <pub-id pub-id-type="doi">10.48550/arXiv.2011.01196</pub-id></citation>
</ref>
<ref id="B58">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Joyce</surname> <given-names>D. W.</given-names></name> <name><surname>Kormilitzin</surname> <given-names>A.</given-names></name> <name><surname>Smith</surname> <given-names>K. A.</given-names></name> <name><surname>Cipriani</surname> <given-names>A.</given-names></name></person-group> (<year>2023</year>). <article-title>Explainable artificial intelligence for mental health through transparency and interpretability for understandability</article-title>. <source>npj Digit. Med</source>. <volume>6</volume>:<fpage>6</fpage>. <pub-id pub-id-type="doi">10.1038/s41746-023-00751-9</pub-id><pub-id pub-id-type="pmid">36653524</pub-id></citation></ref>
<ref id="B59">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Jurafsky</surname> <given-names>D.</given-names></name> <name><surname>Martin</surname> <given-names>J. H.</given-names></name></person-group> (<year>2000</year>). <source>Speech and Language Processing: An Introduction to Natural Language Processing, Computational Linguistics, and Speech Recognition</source>. <publisher-loc>Englewood Cliffs, NJ</publisher-loc>: <publisher-name>Prentice Hall</publisher-name>.</citation>
</ref>
<ref id="B60">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kabir</surname> <given-names>M.</given-names></name> <name><surname>Ahmed</surname> <given-names>T.</given-names></name> <name><surname>Hasan</surname> <given-names>M. B.</given-names></name> <name><surname>Laskar</surname> <given-names>M. T. R.</given-names></name> <name><surname>Joarder</surname> <given-names>T. K.</given-names></name> <name><surname>Mahmud</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Deptweet: a typology for social media texts to detect depression severities</article-title>. <source>Comput. Human Behav</source>. <volume>139</volume>:<fpage>107503</fpage>. <pub-id pub-id-type="doi">10.1016/j.chb.2022.107503</pub-id></citation>
</ref>
<ref id="B61">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kang</surname> <given-names>Y.</given-names></name> <name><surname>Cai</surname> <given-names>Z.</given-names></name> <name><surname>Tan</surname> <given-names>C.-W.</given-names></name> <name><surname>Huang</surname> <given-names>Q.</given-names></name> <name><surname>Liu</surname> <given-names>H.</given-names></name></person-group> (<year>2020</year>). <article-title>Natural language processing (nlp) in management research: a literature review</article-title>. <source>J. Manag. Anal</source>. <volume>7</volume>, <fpage>139</fpage>&#x02013;<lpage>172</lpage>. <pub-id pub-id-type="doi">10.1080/23270012.2020.1756939</pub-id></citation>
</ref>
<ref id="B62">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kawakura</surname> <given-names>S.</given-names></name> <name><surname>Hirafuji</surname> <given-names>M.</given-names></name> <name><surname>Ninomiya</surname> <given-names>S.</given-names></name> <name><surname>Shibasaki</surname> <given-names>R.</given-names></name></person-group> (<year>2022</year>). <article-title>Analyses of diverse agricultural worker data with explainable artificial intelligence: XAI based on shap, lime, and lightgbm</article-title>. <source>Eur. J. Agric. Food Sci</source>. <volume>4</volume>, <fpage>11</fpage>&#x02013;<lpage>19</lpage>. <pub-id pub-id-type="doi">10.24018/ejfood.2022.4.6.348</pub-id></citation>
</ref>
<ref id="B63">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Khan</surname> <given-names>N.</given-names></name> <name><surname>Nauman</surname> <given-names>M.</given-names></name> <name><surname>Almadhor</surname> <given-names>A. S.</given-names></name> <name><surname>Akhtar</surname> <given-names>N.</given-names></name> <name><surname>Alghuried</surname> <given-names>A.</given-names></name> <name><surname>Alhudhaif</surname> <given-names>A.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>Guaranteeing correctness in black-box machine learning: a fusion of explainable AI and formal methods for healthcare decision-making</article-title>. <source>IEEE Access</source> <volume>12</volume>, <fpage>90299</fpage>&#x02013;<lpage>90316</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2024.3420415</pub-id></citation>
</ref>
<ref id="B64">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Khoo</surname> <given-names>L. S.</given-names></name> <name><surname>Lim</surname> <given-names>M. K.</given-names></name> <name><surname>Chong</surname> <given-names>C. Y.</given-names></name> <name><surname>McNaney</surname> <given-names>R.</given-names></name></person-group> (<year>2024</year>). <article-title>Machine learning for multimodal mental health detection: a systematic review of passive sensing approaches</article-title>. <source>Sensors</source> <volume>24</volume>:<fpage>348</fpage>. <pub-id pub-id-type="doi">10.3390/s24020348</pub-id><pub-id pub-id-type="pmid">38257440</pub-id></citation></ref>
<ref id="B65">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kokhlikyan</surname> <given-names>N.</given-names></name> <name><surname>Miglani</surname> <given-names>V.</given-names></name> <name><surname>Martin</surname> <given-names>E.</given-names></name> <name><surname>Wang</surname> <given-names>W.</given-names></name> <name><surname>Alsallakh</surname> <given-names>B.</given-names></name> <name><surname>Reynolds</surname> <given-names>J.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>Captum: a unified and generic model interpretability library for pytorch</article-title>. <source>arXiv [Preprint]</source>. arXiv:<volume>2009</volume>:<fpage>07896</fpage>. <pub-id pub-id-type="doi">10.48550/arXiv.2009:07896</pub-id></citation>
</ref>
<ref id="B66">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>LeCun</surname> <given-names>Y.</given-names></name> <name><surname>Bengio</surname> <given-names>Y.</given-names></name> <name><surname>Hinton</surname> <given-names>G.</given-names></name></person-group> (<year>2015</year>). <article-title>Deep learning</article-title>. <source>Nature</source> <volume>521</volume>, <fpage>436</fpage>&#x02013;<lpage>444</lpage>. <pub-id pub-id-type="doi">10.1038/nature14539</pub-id><pub-id pub-id-type="pmid">26017442</pub-id></citation></ref>
<ref id="B67">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Levy</surname> <given-names>O.</given-names></name> <name><surname>Goldberg</surname> <given-names>Y.</given-names></name></person-group> (<year>2015</year>). <article-title>&#x0201C;Improving distributional similarity with lessons learned from word embeddings,&#x0201D;</article-title> in <source>Transactions of the Association for Computational Linguistics</source> (<publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>MIT Press</publisher-name>), <fpage>211</fpage>&#x02013;<lpage>225</lpage>. <pub-id pub-id-type="doi">10.1162/tacl_a_00134</pub-id></citation>
</ref>
<ref id="B68">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>X.</given-names></name> <name><surname>Jurafsky</surname> <given-names>D.</given-names></name> <name><surname>Liang</surname> <given-names>P.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Analogical reasoning with word vectors: a comparative study of algorithms,&#x0201D;</article-title> in <source>Proceedings of the Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies</source> (<publisher-loc>Santa Fe, NM</publisher-loc>), <fpage>2150</fpage>&#x02013;<lpage>2163</lpage>.</citation>
</ref>
<ref id="B69">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liaw</surname> <given-names>A.</given-names></name> <name><surname>Wiener</surname> <given-names>M.</given-names></name></person-group> (<year>2002</year>). <article-title>Classification and regression by randomforest</article-title>. <source>R News</source> <volume>2</volume>, <fpage>18</fpage>&#x02013;<lpage>22</lpage>.</citation>
</ref>
<ref id="B70">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Manning</surname> <given-names>C. D.</given-names></name> <name><surname>Sch&#x000FC;tze</surname> <given-names>H.</given-names></name></person-group> (<year>1999</year>). <source>Foundations of Statistical Natural Language Processing</source>. MIT Press, Cambridge, MA.</citation>
</ref>
<ref id="B71">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Mohammed</surname> <given-names>M. B.</given-names></name> <name><surname>Abir</surname> <given-names>A. S. M.</given-names></name> <name><surname>Salsabil</surname> <given-names>L.</given-names></name> <name><surname>Shahriar</surname> <given-names>M.</given-names></name> <name><surname>Fahmin</surname> <given-names>A.</given-names></name></person-group> (<year>2021</year>). <article-title>&#x0201C;Depression analysis from social media data in Bangla language: an ensemble approach,&#x0201D;</article-title> in <source>2021 Emerging Technology in Computing, Communication and Electronics (ETCCE)</source> (<publisher-loc>Dhaka</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>1</fpage>&#x02013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1109/ETCCE54784.2021.9689887</pub-id></citation>
</ref>
<ref id="B72">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Nauman</surname> <given-names>M.</given-names></name> <name><surname>Akhtar</surname> <given-names>N.</given-names></name> <name><surname>Alhazmi</surname> <given-names>O. H.</given-names></name> <name><surname>Hameed</surname> <given-names>M.</given-names></name> <name><surname>Ullah</surname> <given-names>H.</given-names></name> <name><surname>Khan</surname> <given-names>N.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Improving the correctness of medical diagnostics based on machine learning with coloured petri nets</article-title>. <source>IEEE Access</source> <volume>9</volume>, <fpage>143434</fpage>&#x02013;<lpage>143447</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2021.3121092</pub-id></citation>
</ref>
<ref id="B73">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Olusegun</surname> <given-names>R.</given-names></name> <name><surname>Oladunni</surname> <given-names>T.</given-names></name> <name><surname>Audu</surname> <given-names>H.</given-names></name> <name><surname>Houkpati</surname> <given-names>Y.</given-names></name> <name><surname>Bengesi</surname> <given-names>S.</given-names></name></person-group> (<year>2023</year>). <article-title>Text mining and emotion classification on monkeypox twitter dataset: a deep learning-natural language processing (NLP) approach</article-title>. <source>IEEE Access</source> <volume>11</volume>, <fpage>49882</fpage>&#x02013;<lpage>49894</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2023.3277868</pub-id></citation>
</ref>
<ref id="B74">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Orabi</surname> <given-names>A. H.</given-names></name> <name><surname>Buddhitha</surname> <given-names>P.</given-names></name> <name><surname>Orabi</surname> <given-names>M. H.</given-names></name> <name><surname>Inkpen</surname> <given-names>D.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x0201C;Deep learning for depression detection of twitter users,&#x0201D;</article-title> in <source>Proceedings of the Fifth Workshop on Computational Linguistics and Clinical Psychology: From Keyboard to Clinic</source>, <fpage>88</fpage>&#x02013;<lpage>97</lpage>.</citation>
</ref>
<ref id="B75">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ossai</surname> <given-names>C. I.</given-names></name> <name><surname>Wickramasinghe</surname> <given-names>N.</given-names></name></person-group> (<year>2023</year>). <article-title>Sentiments prediction and thematic analysis for diabetes mobile apps using embedded deep neural networks and latent Dirichlet allocation</article-title>. <source>Artif. Intell. Med</source>., <volume>138</volume>, <fpage>102509</fpage>. <pub-id pub-id-type="doi">10.1016/j.artmed.2023.102509</pub-id><pub-id pub-id-type="pmid">36990592</pub-id></citation></ref>
<ref id="B76">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Pennington</surname> <given-names>J.</given-names></name> <name><surname>Socher</surname> <given-names>R.</given-names></name> <name><surname>Manning</surname> <given-names>C. D.</given-names></name></person-group> (<year>2014</year>). <article-title>&#x0201C;Glove: global vectors for word representation&#x0201D;</article-title> in <source>Conference on Empirical Methods in Natural Language Processing (EMNLP)</source> (<publisher-loc>Doha</publisher-loc>). <pub-id pub-id-type="doi">10.3115/v1/D14-1162</pub-id></citation>
</ref>
<ref id="B77">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Qasim</surname> <given-names>A.</given-names></name> <name><surname>Mehak</surname> <given-names>G.</given-names></name> <name><surname>Hussain</surname> <given-names>N.</given-names></name> <name><surname>Gelbukh</surname> <given-names>A.</given-names></name> <name><surname>Sidorov</surname> <given-names>G.</given-names></name></person-group> (<year>2025</year>). <article-title>Detection of depression severity in social media text using transformer-based models</article-title>. <source>Information</source> <volume>16</volume>:<fpage>114</fpage>. <pub-id pub-id-type="doi">10.3390/info16020114</pub-id></citation>
</ref>
<ref id="B78">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ram&#x000ED;rez-Cifuentes</surname> <given-names>D.</given-names></name> <name><surname>Largeron</surname> <given-names>C.</given-names></name> <name><surname>Tissier</surname> <given-names>J.</given-names></name> <name><surname>Baeza-Yates</surname> <given-names>R.</given-names></name> <name><surname>Freire</surname> <given-names>A.</given-names></name></person-group> (<year>2021</year>). <article-title>Enhanced word embedding variations for the detection of substance abuse and mental health issues on social media writings</article-title>. <source>IEEE Access</source> <volume>9</volume>, <fpage>130449</fpage>&#x02013;<lpage>130471</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2021.3112102</pub-id></citation>
</ref>
<ref id="B79">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Ribeiro</surname> <given-names>M. T.</given-names></name> <name><surname>Singh</surname> <given-names>S.</given-names></name> <name><surname>Guestrin</surname> <given-names>C.</given-names></name></person-group> (<year>2016a</year>). <article-title>&#x0201C;Model-agnostic interpretability of machine learning,&#x0201D;</article-title> in <source>ICML Workshop on Human Interpretability in Machine Learning (WHI), Volume 31</source> (<publisher-loc>New York, NY</publisher-loc>), <fpage>91</fpage>.</citation>
</ref>
<ref id="B80">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ribeiro</surname> <given-names>M. T.</given-names></name> <name><surname>Singh</surname> <given-names>S.</given-names></name> <name><surname>Guestrin</surname> <given-names>C.</given-names></name></person-group> (<year>2016b</year>). <article-title>&#x0201C;Why should i trust you?&#x0201D; Explaining the predictions of any classifier,&#x0201D;</article-title> in <source>Proceedings of the 22nd ACM SIGKDD International Conference on Knowledge Discovery and Data Mining</source> (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>ACM</publisher-name>), <fpage>1135</fpage>&#x02013;<lpage>1144</lpage>. <pub-id pub-id-type="doi">10.1145/2939672.2939778</pub-id></citation>
</ref>
<ref id="B81">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rizwan</surname> <given-names>M.</given-names></name> <name><surname>Mushtaq</surname> <given-names>M. F.</given-names></name> <name><surname>Akram</surname> <given-names>U.</given-names></name> <name><surname>Mehmood</surname> <given-names>A.</given-names></name> <name><surname>Ashraf</surname> <given-names>I.</given-names></name> <name><surname>Sahelices</surname> <given-names>B.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Depression classification from tweets using small deep transfer learning language models</article-title>. <source>IEEE Access</source> <volume>10</volume>, <fpage>129176</fpage>&#x02013;<lpage>129189</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2022.3223049</pub-id></citation>
</ref>
<ref id="B82">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Sabaneh</surname> <given-names>K.</given-names></name> <name><surname>Salameh</surname> <given-names>M. A.</given-names></name> <name><surname>Khaleel</surname> <given-names>F.</given-names></name> <name><surname>Herzallah</surname> <given-names>M. M.</given-names></name> <name><surname>Natsheh</surname> <given-names>J. Y.</given-names></name> <name><surname>Maree</surname> <given-names>M.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>&#x0201C;Early risk prediction of depression based on social media posts in Arabic,&#x0201D;</article-title> in <source>2023 IEEE 35th International Conference on Tools with Artificial Intelligence (ICTAI)</source> (<publisher-loc>Atlanta, GA</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>595</fpage>&#x02013;<lpage>602</lpage>. <pub-id pub-id-type="doi">10.1109/ICTAI59109.2023.00094</pub-id></citation>
</ref>
<ref id="B83">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Saha</surname> <given-names>T.</given-names></name> <name><surname>Reddy</surname> <given-names>S. M.</given-names></name> <name><surname>Saha</surname> <given-names>S.</given-names></name> <name><surname>Bhattacharyya</surname> <given-names>P.</given-names></name></person-group> (<year>2022</year>). <article-title>Mental health disorder identification from motivational conversations</article-title>. <source>IEEE Trans. Comput. Soc. Syst</source>. <volume>D10</volume>, <fpage>1130</fpage>&#x02013;<lpage>1139</lpage>. <pub-id pub-id-type="doi">10.1109/TCSS.2022.3143763</pub-id></citation>
</ref>
<ref id="B84">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Salton</surname> <given-names>G.</given-names></name> <name><surname>Buckley</surname> <given-names>C.</given-names></name></person-group> (<year>1988</year>). <article-title>Term-weighting approaches in automatic text retrieval</article-title>. <source>Inform. Process. Manag</source>. <volume>24</volume>, <fpage>513</fpage>&#x02013;<lpage>523</lpage>. <pub-id pub-id-type="doi">10.1016/0306-4573(88)90021-0</pub-id></citation>
</ref>
<ref id="B85">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Santhosh Baboo</surname> <given-names>S.</given-names></name> <name><surname>Amirthapriya</surname> <given-names>M.</given-names></name></person-group> (<year>2022</year>). <article-title>Comparison of machine learning techniques on twitter emotions classification</article-title>. <source>SN Comput. Sci</source>. <volume>3</volume>, <fpage>1</fpage>&#x02013;<lpage>8</lpage>. <pub-id pub-id-type="doi">10.1007/s42979-021-00889-x</pub-id></citation>
</ref>
<ref id="B86">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Santos</surname> <given-names>W.</given-names></name> <name><surname>Yoon</surname> <given-names>S.</given-names></name> <name><surname>Paraboni</surname> <given-names>I.</given-names></name></person-group> (<year>2023</year>). <article-title>Mental health prediction from social media text using mixture of experts</article-title>. <source>IEEE Latin Am. Trans</source>. <volume>21</volume>, <fpage>723</fpage>&#x02013;<lpage>729</lpage>. <pub-id pub-id-type="doi">10.1109/TLA.2023.10172137</pub-id></citation>
</ref>
<ref id="B87">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Saxena</surname> <given-names>C.</given-names></name> <name><surname>Garg</surname> <given-names>M.</given-names></name> <name><surname>Ansari</surname> <given-names>G.</given-names></name></person-group> (<year>2022</year>). &#x0201C;Explainable causal analysis of mental health on social media data,&#x0201D; <italic>International Conference on Neural Information Processing</italic> (<publisher-loc>Cham</publisher-loc>: <publisher-name>Springer</publisher-name>), <fpage>172</fpage>&#x02013;<lpage>183</lpage>. <pub-id pub-id-type="doi">10.1007/978-3-031-30108-7_15</pub-id></citation>
</ref>
<ref id="B88">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schmidhuber</surname> <given-names>J.</given-names></name></person-group> (<year>2015</year>). <article-title>Deep learning in neural networks: an overview</article-title>. <source>Neural Netw</source>. <volume>61</volume>, <fpage>85</fpage>&#x02013;<lpage>117</lpage>. <pub-id pub-id-type="doi">10.1016/j.neunet.2014.09.003</pub-id><pub-id pub-id-type="pmid">25462637</pub-id></citation></ref>
<ref id="B89">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Sch&#x000F6;lkopf</surname> <given-names>B.</given-names></name> <name><surname>Smola</surname> <given-names>A. J.</given-names></name></person-group> (<year>2002</year>). <source>Learning with Kernels: Support Vector Machines, Regularization, Optimization, and Beyond</source>. <publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>MIT press</publisher-name>. <pub-id pub-id-type="doi">10.7551/mitpress/4175.001.0001</pub-id></citation>
</ref>
<ref id="B90">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shah</surname> <given-names>S. M.</given-names></name> <name><surname>Gillani</surname> <given-names>S. A.</given-names></name> <name><surname>Baig</surname> <given-names>M. S. A.</given-names></name> <name><surname>Saleem</surname> <given-names>M. A.</given-names></name> <name><surname>Siddiqui</surname> <given-names>M. H.</given-names></name></person-group> (<year>2025</year>). <article-title>Advancing depression detection on social media platforms through fine-tuned large language models</article-title>. <source>Online Soc. Netw. Media</source> <volume>46</volume>:<fpage>100311</fpage>. <pub-id pub-id-type="doi">10.1016/j.osnem.2025.100311</pub-id><pub-id pub-id-type="pmid">40619512</pub-id></citation></ref>
<ref id="B91">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>S&#x000F8;gaard</surname> <given-names>A.</given-names></name></person-group> (<year>2021</year>). <source>Explainable Natural Language Processing</source>. <publisher-loc>San Rafael, CA</publisher-loc>: <publisher-name>Morgan &#x00026; Claypool Publishers</publisher-name>.</citation>
</ref>
<ref id="B92">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Soldaini</surname> <given-names>L.</given-names></name> <name><surname>Goharian</surname> <given-names>N.</given-names></name></person-group> (<year>2016</year>). <article-title>&#x0201C;Quickumls: a fast, unsupervised approach for medical concept extraction,&#x0201D;</article-title> in <source>Proceedings of the AMIA Annual Symposium</source> (<publisher-loc>Washington, DC</publisher-loc>: <publisher-name>American Medical Informatics Association</publisher-name>), <fpage>1216</fpage>&#x02013;<lpage>1225</lpage>.</citation>
</ref>
<ref id="B93">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Steinkamp</surname> <given-names>J.</given-names></name> <name><surname>Cook</surname> <given-names>T. S.</given-names></name></person-group> (<year>2021a</year>). <article-title>Basic artificial intelligence techniques: natural language processing of radiology reports</article-title>. <source>Radiol. Clin</source>. <volume>59</volume>, <fpage>919</fpage>&#x02013;<lpage>931</lpage>. <pub-id pub-id-type="doi">10.1016/j.rcl.2021.06.003</pub-id><pub-id pub-id-type="pmid">34689877</pub-id></citation></ref>
<ref id="B94">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Steinkamp</surname> <given-names>J. M.</given-names></name> <name><surname>Cook</surname> <given-names>T. B.</given-names></name></person-group> (<year>2021b</year>). <article-title>Applications of natural language processing for mental health research on social media</article-title>. <source>Int. J. Med. Inform</source>. <volume>148</volume>:<fpage>104399</fpage>.</citation>
</ref>
<ref id="B95">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Sundararajan</surname> <given-names>M.</given-names></name> <name><surname>Taly</surname> <given-names>A.</given-names></name> <name><surname>Yan</surname> <given-names>Q.</given-names></name></person-group> (<year>2017</year>). <article-title>&#x0201C;Axiomatic attribution for deep networks,&#x0201D;</article-title> in <source>Proceedings of the 34th International Conference on Machine Learning (ICML)</source> (<publisher-loc>Sydney, NSW</publisher-loc>: <publisher-name>PMLR</publisher-name>), <fpage>3319</fpage>&#x02013;<lpage>3328</lpage>.</citation>
</ref>
<ref id="B96">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tadesse</surname> <given-names>M. M.</given-names></name> <name><surname>Lin</surname> <given-names>H.</given-names></name> <name><surname>Xu</surname> <given-names>B.</given-names></name> <name><surname>Yang</surname> <given-names>L.</given-names></name></person-group> (<year>2019</year>). <article-title>Detection of depression-related posts in reddit social media forum</article-title>. <source>IEEE Access</source> <volume>7</volume>, <fpage>44883</fpage>&#x02013;<lpage>44893</lpage>. <pub-id pub-id-type="doi">10.1109/ACCESS.2019.2909180</pub-id></citation>
</ref>
<ref id="B97">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Thornicroft</surname> <given-names>G.</given-names></name> <name><surname>Sunkel</surname> <given-names>C.</given-names></name> <name><surname>Aliev</surname> <given-names>A. A.</given-names></name> <name><surname>Baker</surname> <given-names>S.</given-names></name> <name><surname>Brohan</surname> <given-names>E.</given-names></name> <name><surname>El Chammay</surname> <given-names>R.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>The lancet commission on ending stigma and discrimination in mental health</article-title>. <source>Lancet</source> <volume>400</volume>, <fpage>1438</fpage>&#x02013;<lpage>1480</lpage>. <pub-id pub-id-type="doi">10.1016/S0140-6736(22)01470-2</pub-id><pub-id pub-id-type="pmid">36223799</pub-id></citation></ref>
<ref id="B98">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tlachac</surname> <given-names>M.</given-names></name> <name><surname>Rundensteiner</surname> <given-names>E.</given-names></name></person-group> (<year>2020</year>). <article-title>Screening for depression with retrospectively harvested private versus public text</article-title>. <source>IEEE J. Biomed. Health Inform</source>. <volume>24</volume>, <fpage>3326</fpage>&#x02013;<lpage>3332</lpage>. <pub-id pub-id-type="doi">10.1109/JBHI.2020.2983035</pub-id><pub-id pub-id-type="pmid">32224470</pub-id></citation></ref>
<ref id="B99">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tong</surname> <given-names>L.</given-names></name> <name><surname>Liu</surname> <given-names>Z.</given-names></name> <name><surname>Jiang</surname> <given-names>Z.</given-names></name> <name><surname>Zhou</surname> <given-names>F.</given-names></name> <name><surname>Chen</surname> <given-names>L.</given-names></name> <name><surname>Lyu</surname> <given-names>J.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Cost-sensitive boosting pruning trees for depression detection on Twitter</article-title>. <source>IEEE Trans. Affect. Comput</source>. <volume>14</volume>, <fpage>1898</fpage>&#x02013;<lpage>1911</lpage>. <pub-id pub-id-type="doi">10.1109/TAFFC.2022.3145634</pub-id></citation>
</ref>
<ref id="B100">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Vapnik</surname> <given-names>V. N.</given-names></name></person-group> (<year>1998</year>). <source>Statistical Learning Theory</source>. <publisher-loc>Hoboken, NJ</publisher-loc>: <publisher-name>Wiley</publisher-name>.</citation>
</ref>
<ref id="B101">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Velupillai</surname> <given-names>S.</given-names></name> <name><surname>Suominen</surname> <given-names>H.</given-names></name> <name><surname>Liakata</surname> <given-names>M.</given-names></name> <name><surname>Roberts</surname> <given-names>A.</given-names></name> <name><surname>Shah</surname> <given-names>A. D.</given-names></name> <name><surname>Morley</surname> <given-names>K.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Using clinical natural language processing for health outcomes research: overview and actionable suggestions for future advances</article-title>. <source>J. Biomed. Inform</source>. <volume>88</volume>, <fpage>11</fpage>&#x02013;<lpage>19</lpage>. <pub-id pub-id-type="doi">10.1016/j.jbi.2018.10.005</pub-id><pub-id pub-id-type="pmid">30368002</pub-id></citation></ref>
<ref id="B102">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wani</surname> <given-names>M. A.</given-names></name> <name><surname>ELAffendi</surname> <given-names>M. A.</given-names></name> <name><surname>Shakil</surname> <given-names>K. A.</given-names></name> <name><surname>Imran</surname> <given-names>A. S.</given-names></name> <name><surname>Abd El-Latif</surname> <given-names>A. A.</given-names></name></person-group> (<year>2022</year>). <article-title>Depression screening in humans with AI and deep learning techniques</article-title>. <source>IEEE Trans. Comput. Soc. Syst</source>. <volume>10</volume>, <fpage>2074</fpage>&#x02013;<lpage>2089</lpage>. <pub-id pub-id-type="doi">10.1109/TCSS.2022.3200213</pub-id></citation>
</ref>
<ref id="B103">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Weerts</surname> <given-names>H. J.</given-names></name> <name><surname>van Ipenburg</surname> <given-names>W.</given-names></name> <name><surname>Pechenizkiy</surname> <given-names>M.</given-names></name></person-group> (<year>2019</year>). <article-title>A human-grounded evaluation of shap for alert processing</article-title>. <source>arXiv [Preprint]</source>. arXiv:1907.03324. <pub-id pub-id-type="doi">10.48550/arXiv.1907.03324</pub-id></citation>
</ref>
<ref id="B104">
<citation citation-type="journal"><person-group person-group-type="author"><collab>World Health Organization</collab></person-group> (<year>2021</year>). Guidance on Community Mental Health Services: Promoting Person-Centred and Rights-Based Approaches. Geneva: World Health Organization.</citation>
</ref>
<ref id="B105">
<citation citation-type="journal"><person-group person-group-type="author"><collab>World Health Organization Regional Office for the Eastern Mediterranean.</collab></person-group> (<year>2019</year>). <source>Mental Health Atlas 2017: Resources for Mental Health in the Eastern Mediterranean Region</source>.<pub-id pub-id-type="pmid">28493259</pub-id></citation></ref>
<ref id="B106">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xiang</surname> <given-names>L.</given-names></name></person-group> (<year>2022</year>). <article-title>Application of an improved tf-idf method in literary text classification</article-title>. <source>Adv. Multimed</source>. <volume>2022</volume>:<fpage>9285324</fpage>. <pub-id pub-id-type="doi">10.1155/2022/9285324</pub-id></citation>
</ref>
<ref id="B107">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Xu</surname> <given-names>Z.</given-names></name> <name><surname>Chen</surname> <given-names>M.</given-names></name> <name><surname>Weinberger</surname> <given-names>K. Q.</given-names></name> <name><surname>Sha</surname> <given-names>F.</given-names></name></person-group> (<year>2013</year>). <article-title>&#x0201C;An alternative text representation to TF-IDF and bag-of-words,&#x0201D;</article-title> in <source>Proceedings of ICML</source> (<italic>arXiv</italic> [Preprint]. arXiv:1301.6770). <pub-id pub-id-type="doi">10.48550/arXiv.1301.6770</pub-id><pub-id pub-id-type="pmid">27001195</pub-id></citation></ref>
<ref id="B108">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yazdavar</surname> <given-names>A. H.</given-names></name> <name><surname>Mahdavinejad</surname> <given-names>M. S.</given-names></name> <name><surname>Bajaj</surname> <given-names>G.</given-names></name> <name><surname>Romine</surname> <given-names>W.</given-names></name> <name><surname>Sheth</surname> <given-names>A.</given-names></name> <name><surname>Monadjemi</surname> <given-names>A. H.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>Multimodal mental health analysis in social media</article-title>. <source>PLoS ONE</source> <volume>15</volume>:<fpage>e0226248</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0226248</pub-id><pub-id pub-id-type="pmid">32275658</pub-id></citation></ref>
<ref id="B109">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>T.</given-names></name> <name><surname>Schoene</surname> <given-names>A. M.</given-names></name> <name><surname>Ji</surname> <given-names>S.</given-names></name> <name><surname>Ananiadou</surname> <given-names>S.</given-names></name></person-group> (<year>2022</year>). <article-title>Natural language processing applied to mental illness detection: a narrative review</article-title>. <source>NPJ Digit. Med</source>. <volume>5</volume>, <fpage>1</fpage>&#x02013;<lpage>13</lpage>. <pub-id pub-id-type="doi">10.1038/s41746-022-00589-7</pub-id><pub-id pub-id-type="pmid">35396451</pub-id></citation></ref>
<ref id="B110">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>W.</given-names></name> <name><surname>Zhang</surname> <given-names>S.</given-names></name> <name><surname>Zhou</surname> <given-names>J.</given-names></name></person-group> (<year>2017</year>). <article-title>Distributed xgboost with column block splitting</article-title>. <source>arXiv [Preprint]</source>. arXiv:1708.05721. <pub-id pub-id-type="doi">10.48550/arXiv.1708.05721</pub-id></citation>
</ref>
<ref id="B111">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Weng</surname> <given-names>Y.</given-names></name> <name><surname>Lund</surname> <given-names>J.</given-names></name></person-group> (<year>2022</year>). <article-title>Applications of explainable artificial intelligence in diagnosis and surgery</article-title>. <source>Diagnostics</source> <volume>12</volume>:<fpage>237</fpage>. <pub-id pub-id-type="doi">10.3390/diagnostics12020237</pub-id><pub-id pub-id-type="pmid">35204328</pub-id></citation></ref>
</ref-list>
</back>
</article>