<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD JATS (Z39.96) Journal Publishing DTD v1.3 20210610//EN" "JATS-journalpublishing1-3-mathml3.dtd">
<article article-type="research-article" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:ali="http://www.niso.org/schemas/ali/1.0/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" dtd-version="1.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Artif. Intell.</journal-id>
<journal-title-group>
<journal-title>Frontiers in Artificial Intelligence</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Artif. Intell.</abbrev-journal-title>
</journal-title-group>
<issn pub-type="epub">2624-8212</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/frai.2025.1537432</article-id><article-version article-version-type="Version of Record" vocab="NISO-RP-8-2008"/>
<article-categories>
<subj-group subj-group-type="heading"><subject>Original Research</subject></subj-group>
</article-categories>
<title-group>
<article-title>Explainable detection: a transformer-based language modeling approach for Bengali news title classification with comparative explainability analysis using ML and DL</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Naeen</surname>
<given-names>Md. Julkar</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/3091291"/>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="conceptualization" vocab-term-identifier="https://credit.niso.org/contributor-roles/conceptualization/">Conceptualization</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Data curation" vocab-term-identifier="https://credit.niso.org/contributor-roles/data-curation/">Data curation</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Formal analysis" vocab-term-identifier="https://credit.niso.org/contributor-roles/formal-analysis/">Formal analysis</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="methodology" vocab-term-identifier="https://credit.niso.org/contributor-roles/methodology/">Methodology</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; original draft" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-original-draft/">Writing &#x2013; original draft</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &amp; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-review-editing/">Writing &#x2013; review &#x0026; editing</role>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Das</surname>
<given-names>Sourav Kumar</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2910210"/>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="conceptualization" vocab-term-identifier="https://credit.niso.org/contributor-roles/conceptualization/">Conceptualization</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Data curation" vocab-term-identifier="https://credit.niso.org/contributor-roles/data-curation/">Data curation</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Formal analysis" vocab-term-identifier="https://credit.niso.org/contributor-roles/formal-analysis/">Formal analysis</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="methodology" vocab-term-identifier="https://credit.niso.org/contributor-roles/methodology/">Methodology</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="resources" vocab-term-identifier="https://credit.niso.org/contributor-roles/resources/">Resources</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; original draft" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-original-draft/">Writing &#x2013; original draft</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &amp; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-review-editing/">Writing &#x2013; review &#x0026; editing</role>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Jisan</surname>
<given-names>Sakib Alam</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/3233971"/>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="conceptualization" vocab-term-identifier="https://credit.niso.org/contributor-roles/conceptualization/">Conceptualization</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; original draft" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-original-draft/">Writing &#x2013; original draft</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &amp; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-review-editing/">Writing &#x2013; review &#x0026; editing</role>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Khushbu</surname>
<given-names>Sharun Akter</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="supervision" vocab-term-identifier="https://credit.niso.org/contributor-roles/supervision/">Supervision</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; original draft" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-original-draft/">Writing &#x2013; original draft</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &amp; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-review-editing/">Writing &#x2013; review &#x0026; editing</role>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Saha</surname>
<given-names>Noyon Chandra</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="methodology" vocab-term-identifier="https://credit.niso.org/contributor-roles/methodology/">Methodology</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &amp; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-review-editing/">Writing &#x2013; review &#x0026; editing</role>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Ohidujjaman</surname>
</name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2745274"/>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="supervision" vocab-term-identifier="https://credit.niso.org/contributor-roles/supervision/">Supervision</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &amp; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-review-editing/">Writing &#x2013; review &#x0026; editing</role>
</contrib>
</contrib-group>
<aff id="aff1"><label>1</label><institution>Department of Computer Science and Engineering, Daffodil International University</institution>, <city>Dhaka</city>, <country country="bd">Bangladesh</country></aff>
<aff id="aff2"><label>2</label><institution>Department of Computer Science and Engineering, United International University</institution>, <city>Dhaka</city>, <country country="bd">Bangladesh</country></aff>
<author-notes><corresp id="c001"><label>&#x002A;</label>Correspondence: Sourav Kumar Das, <email xlink:href="mailto:sourav15-4588@diu.edu.bd">sourav15-4588@diu.edu.bd</email></corresp></author-notes>
<pub-date publication-format="electronic" date-type="pub" iso-8601-date="2025-11-06">
<day>06</day>
<month>11</month>
<year>2025</year>
</pub-date>
<pub-date publication-format="electronic" date-type="collection">
<year>2025</year>
</pub-date>
<volume>8</volume>
<elocation-id>1537432</elocation-id>
<history>
<date date-type="received">
<day>30</day>
<month>11</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>29</day>
<month>09</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2025 Naeen, Das, Jisan, Khushbu, Saha and Ohidujjaman.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Naeen, Das, Jisan, Khushbu, Saha and Ohidujjaman</copyright-holder>
<license><ali:license_ref start_date="2025-11-06">https://creativecommons.org/licenses/by/4.0/</ali:license_ref>
<license-p>This is an open-access article distributed under the terms of the <ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">Creative Commons Attribution License (CC BY)</ext-link>. The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</license-p>
</license>
</permissions>
<abstract>
<p>Classifying scattered Bengali text is the primary focus of this study, with an emphasis on explainability in Natural Language Processing (NLP) for low-resource languages. We employed supervised Machine Learning (ML) models as a baseline and compared their performance with Long Short-Term Memory (LSTM) networks from the deep learning domain. Subsequently, we implemented transformer models designed for sequential learning. To prepare the dataset, we collected recent Bengali news articles online and performed extensive feature engineering. Given the inherent noise in Bengali datasets, significant preprocessing was required. Among the models tested, XLM-RoBERTa Base achieved the highest accuracy 0.91. Furthermore, we integrated explainable AI techniques to interpret the model&#x2019;s predictions, enhancing transparency and fostering trust in the classification outcomes. Additionally, we employed LIME (Local Interpretable Model-agnostic Explanations) to identify key features and the most weighted words responsible for classifying news titles, which validated the accuracy of Bengali news classification results. This study underscores the potential of deep learning models in advancing text classification for the Bengali language and emphasizes the critical role of explainability in AI-driven solutions.</p>
</abstract>
<kwd-group>
<kwd>transformer model</kwd>
<kwd>long short-term memory</kwd>
<kwd>Bengali news titles</kwd>
<kwd>classification</kwd>
<kwd>LIME</kwd>
<kwd>explainable AI</kwd>
<kwd>machine learning</kwd>
<kwd>deep learning</kwd>
</kwd-group><funding-group><award-group id="gs1"><funding-source id="sp1"><institution-wrap><institution>Institute for Advanced Research Publication Grant of United International University</institution></institution-wrap></funding-source><award-id rid="sp1">IAR-2025-Pub-050</award-id></award-group><funding-statement>The author(s) declare that financial support was received for the research and/or publication of this article. This research was funded by the Institute for Advanced Research, United International University (UIU), Ref. No.: IAR-2025-Pub-050.</funding-statement></funding-group><counts>
<fig-count count="8"/>
<table-count count="11"/>
<equation-count count="12"/>
<ref-count count="59"/>
<page-count count="15"/>
<word-count count="10540"/>
</counts>
<custom-meta-group>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Natural Language Processing</meta-value>
</custom-meta>
</custom-meta-group>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec1">
<label>1</label>
<title>Introduction</title>
<p>Information is the most significant asset in the modern world. People utilize various platforms to access information, with newspapers being one of the most common and accessible sources. Newspapers offer a wealth of information on diverse topics at an affordable price, making knowledge accessible to everyone. They enrich readers&#x2019; understanding and provide insight into domestic and international current events. Newspaper has become quite easy in today&#x2019;s world. Humans can easily comprehend news headlines and their underlying meanings due to their familiarity with the language and context. However, this task poses significant challenges for machines, particularly when processing text in Bengali. In the Bengali language, many words have multiple meanings that vary depending on their context and usage, making it difficult for machines to interpret their intended meaning accurately. Additionally, some news headlines are lengthy, further complicating the extraction of semantic information from these complex sentences. To address this issue, it is essential to train models capable of understanding the contextual meaning of sentences. Transformer-based models are particularly well-suited for this task, as they leverage a deeply bidirectional architecture, enabling them to capture the contextual relationships within a sentence. Consequently, transformers are a robust choice for deriving the semantic meaning of words and sentences, offering a significant advantage in tasks involving natural language understanding. Bengali, the national language of Bangladesh, is spoken by approximately 300 million people worldwide, drawing significant attention in the field of Natural Language Processing (NLP) (<xref ref-type="bibr" rid="ref29">Hossain et al., 2020a</xref>). Recent research on Bengali text has been extensive. In response, we aimed to innovate in Bengali news article classification. We gathered raw data from various newspapers, balanced it for better performance, processed it, and applied machine learning and deep learning models, along with explainable AI. Our goal was to classify articles based on their titles. Our model can predict the category from headlines of varying lengths, effectively handling Bengali words with multiple meanings depending on context. This capability enhances our research&#x2019;s accuracy. If successful, our study could spark further interest among NLP researchers. Classifying Bengali newspaper articles is challenging due to certain linguistic complexities. However, overcoming these challenges could yield promising results, as deep learning and NLP provide optimal solutions for text classification problems.</p>
<p>Bengali text classification is quite popular nowadays. In the newspaper, people give different opinions about national, international, politics, sports, etc. Our work is related to identifying different classes from titles. Sentiment analysis is a prominent aspect of NLP research, emphasizing the importance of identifying words that convey positive and negative meanings (<xref ref-type="bibr" rid="ref48">Roy et al., 2023</xref>). Additionally, fake news, prevalent even in newspapers, can mislead people and obscure the truth (<xref ref-type="bibr" rid="ref21">Fouad et al., 2022</xref>). Most algorithms struggle with plain text, making word embedding understanding essential (<xref ref-type="bibr" rid="ref57">Wadud et al., 2022</xref>). The Transformer is a recent neural network model, well-supported for English but lacking resources for Bengali text classification (<xref ref-type="bibr" rid="ref4">Alam et al., 2020</xref>). Fake news detection is also prevalent in other languages (<xref ref-type="bibr" rid="ref16">Das et al., 2023a</xref>). Despite the significant resource gap for the Bengali language in NLP, some researchers have managed to categorize Bengali sentences into different forms (<xref ref-type="bibr" rid="ref17">Das et al., 2023b</xref>). Sentiment analysis remains crucial, successfully detecting emotions in sentences (<xref ref-type="bibr" rid="ref9">Bhowmik et al., 2021</xref>). For emotions like anger, disgust, fear, joy, sadness, and surprise in Bengali, researchers have proposed an interesting transformer-based method (<xref ref-type="bibr" rid="ref53">Sourav et al., 2022</xref>). Cyberbullying is a common issue today, prompting NLP researchers to explore prevention methods (<xref ref-type="bibr" rid="ref1">Ahmed et al., 2021</xref>). Malicious activities targeting government security also pose a significant problem (<xref ref-type="bibr" rid="ref6">Aslam et al., 2022</xref>). LIME is an effective tool for explaining black-box machine learning models in various fields (<xref ref-type="bibr" rid="ref56">Venkatsubramaniam and Baruah, 2022</xref>).</p>
<p>Our findings reveal a substantial body of research on the Bengali language. However, comparatively limited work focuses on classifying and identifying the semantic meaning of words or sentences within the contextually rich and often ambiguous structure of Bengali. Many Bengali sentences carry multiple meanings depending on their situational and contextual usage, posing significant challenges for machines in accurately discerning their underlying meaning. This research seeks to address these challenges. Several machine learning and deep learning models, including a Long Short-Term Memory (LSTM) network, were employed in this study. While these models demonstrated strong performance in classifying the dataset, they fell short in capturing words and sentences&#x2019; deeper, contextual semantics. LSTM and traditional machine learning models struggle to understand nuanced meanings that depend heavily on context. In contrast, transformer-based models named XLM-RoBERTa base (<xref ref-type="bibr" rid="ref12">Conneau, 2019</xref>) and Multilingual BERT (<xref ref-type="bibr" rid="ref34">Kenton and Toutanova, 2019</xref>) outperformed these approaches. Due to their deeply bidirectional architecture and superior capability to learn contextual and semantic nuances, transformers are more effective in understanding the true meaning of sentences. This makes them a more suitable choice for processing the Bengali language, where semantic ambiguity is prevalent.</p>
<p>This study aims to explore the following research questions:</p>
<list list-type="order">
<list-item><p>How can Bengali news headlines be accurately classified despite the contextual ambiguity and multiple meanings of Bengali words?</p></list-item>
<list-item><p>To what extent can transformer-based models, such as XLM-RoBERTa and Multilingual BERT, outperform traditional machine learning and deep learning models (e.g., LSTM) in classifying Bengali newspaper headlines?</p></list-item>
<list-item><p>Can Explainable AI techniques, such as LIME, effectively reveal which parts of Bengali headlines influence model decisions, thereby improving interpretability?</p></list-item>
</list>
<p>We contributed to our dataset by creating our own properly annotated data on news lines collected from various sources. We also contributed to the comparison of conventional approaches, deep learning, and transformer models. We customized the squeeze and attention blocks in the transformer model to achieve smaller weights, enhancing efficiency and producing a low-loss graph. After evaluating the model&#x2019;s performance, we utilized Explainable AI (XAI) to gain deeper insights into Bengali word interpretations within the hidden layers. This helped identify which words contributed more to the model and which were processed for the next iteration. We applied the LIME technique, which revealed that headline-related words carried higher weights during text learning.</p>
</sec>
<sec id="sec2">
<label>2</label>
<title>Related work</title>
<p>Until recently, there was limited research on the Bengali language in fields like machine learning, deep learning, and NLP. There has been some work done in this area, but it is still not very advanced and is not very common. Several studies have employed various BERT models, but none have incorporated Explainable AI (XAI) techniques. Despite using BERT, some report lower accuracy than our model. Our approach, integrating both BERT and XAI, yields improved performance. In earlier studies, models like Random Forest, Multinomial Naive Bayes, and LSTM were used in similar ways. BERT made these methods more effective and efficient.</p>
<p><xref ref-type="bibr" rid="ref40">Maisha et al. (2021)</xref> did sentiment analysis of Bengali newspapers by implementing supervised machine learning algorithms. Some techniques were combined for the class. From all the six models, Random Forest provided the best accuracy of 99%. <xref ref-type="bibr" rid="ref8">Bhowmik et al. (2022)</xref> did the sentiment analysis with an extended lexicon dictionary and deep learning, and the highest accuracy was in the BERT-LSTM model. As well, sentiment analysis was done by <xref ref-type="bibr" rid="ref25">Hassan et al. (2022)</xref> for Bengali conversation, and support vector machine (SVM) gave the best accuracy, which is 85.59%. <xref ref-type="bibr" rid="ref43">Prottasha et al. (2022)</xref> also did sentiment analysis on behalf of BERT-based supervised fine-tuning. Word embedding techniques like Word2Vec, GloVe, and fastText are used. CNN-BiLSTM provided the highest accuracy of 94.15%. After that, <xref ref-type="bibr" rid="ref35">Keya et al. (2022)</xref> also used BERT to classify fake news and created the AugFake-BERT model. To implement the model, more than 50,000 data points are used. However, the proposed model provided an accuracy of 92.45%, and all the other scores are utilized to evaluate the performance. By using BERT, <xref ref-type="bibr" rid="ref37">Kowsher et al. (2022)</xref> created the Bengali-BERT model for language understanding and transfer learning. Bengali-BERT performed better than the other models, with 97.03% accuracy. Then, <xref ref-type="bibr" rid="ref28">Hossain et al. (2020b)</xref> categorized Bengali news headlines with deep learning models. For classification, two models are used: LSTM and GRU. Both models provided almost the same accuracy but with a little difference. GRU shows the highest accuracy of 87.74%. <xref ref-type="bibr" rid="ref30">Hossain et al. (2020c)</xref> classified Bengali news using dissimilar machine learning-based baseline approaches and deep learning models. A total of 3,000 data were used to implement the models. SVC, LSVC, Random Forest, Linear Regression, Naive Bayes, CNN, and BiLSTM models are applied for the classification, and the highest accuracy of 93.43% came from the CNN model. <xref ref-type="bibr" rid="ref19">Dhar and Morshed (2022)</xref> analyzed Bengali crime news categorization with the help of machine learning models. From different newspapers, a total of 3,500 data were collected for the implementation. After training the data, the proposed model shows a test accuracy of 87%. Subsequently, <xref ref-type="bibr" rid="ref58">Yeasmin et al. (2021)</xref> proposed a topic about Bengali news classification using ML and Neural Network models. A total of two datasets were used for the process. One of the datasets was collected from the Bengali newspapers, and the other was collected from Kaggle. One neutral network model provided the highest accuracy of 92.63% for dataset 1. For dataset 2, a neural network model again offered the highest accuracy of 95.50%. Applying CNN, RNN, and other deep learning models might give better outcomes. <xref ref-type="bibr" rid="ref59">Zhang (2021)</xref> also researched the application of deep learning in news text classification on different datasets. However, the models were CNN, MLP, LSTM, and some hybrid models were used, and one hybrid model outnumbered all other models and displayed 94.82% accuracy. Similarly, <xref ref-type="bibr" rid="ref45">Ramdhani et al. (2020)</xref> used convolutional neural networks (CNN) to classify Indonesian news. CNN has the best accuracy of 90.74%, with a value loss of 29.05%. <xref ref-type="bibr" rid="ref49">Saigal and Khanna (2020)</xref> applied SVM-based classifiers to classify the category of news. Some ML and hybrid deep learning models were used in the research, and a hybrid model named LS-TWSVM showed the highest accuracy, which was 98.21% on a specific dataset named the Reuters dataset. The dataset was collected from UCI News datasets like Reuters and 20 Newsgroups. After that, <xref ref-type="bibr" rid="ref41">Mridha et al. (2021)</xref> created the L-Boost model, which can identify abusive words from social media posts in Bengali. ML model AdaBoost and DL model LSTM are combined with a transformer (BERT). The model reached 95.11% accuracy, which is the highest among all the ML and DL models. <xref ref-type="bibr" rid="ref52">Sen et al. (2022)</xref> processed Bengali natural language for comprehensive analysis. Classical, machine learning, and deep learning applied on the study. A total of 75 BNLP research papers were studied and categorized into 11 categories for the research. Here we examined recent works that employed similar approaches, including machine learning and deep learning models such as Multinomial Naive Bayes, SVC, LSTM, as well as techniques like TF-IDF, BERT, and others.</p>
<p>At present, a lot of research is going on in the Bengali language. Using machine learning and deep learning models, many studies are ongoing. Just like that, <xref ref-type="bibr" rid="ref23">Hasan et al. (2023a)</xref> classified Bengali newspaper headlines by using LSTM, Bi-LSTM, and Bi-GRU models. Almost 10,000 data points were classified into six categories to achieve the expected result. From those three deep learning models, the Bi-LSTM provides the highest train and test accuracy of 97.96 and 77.91%, respectively. <xref ref-type="bibr" rid="ref3">Al Mahmud et al. (2023)</xref> similarly proposed an approach to classifying Bengali news by using machine learning and deep learning models. Applied models are SVC, Random Forest, LSVC, LSTM, and GRU. The approach technique provided an accuracy of 95.45%, which is the highest among all the algorithms. <xref ref-type="bibr" rid="ref31">Hussain et al. (2023)</xref> also did some comparison analysis of Bengali news article classification using some ML models. TF-IDF and count vectorizer were used for the feature extraction process. SVM and LR algorithms were applied, and SVM provided the highest accuracy of 84%. There were 20 categories, and 12.5&#x202F;K labeled news articles were used. Adding more categories and applying more ML and deep learning might give a more optimal result. After that, <xref ref-type="bibr" rid="ref39">Mahmud et al. (2023)</xref> evaluated news by using Natural Language Processing (NLP) and Human Expert Opinion. A total of three NLP models were applied for training and testing. The Bengali-Bertbase model provided the highest testing accuracy of 84.99%, and it increased after the 9th parameter, where it achieved 93.80% of testing accuracy. Then, <xref ref-type="bibr" rid="ref24">Hasan et al. (2023b)</xref> did sentiment analysis and natural language processing (NLP) using transformers for Russia-Ukraine war-based comments in Bengali. The applied models for the analysis are mBERT, Distil-mBERT, BengaliBERT, XML-R(base), XML-R (large), and Bi-LSTM. The BengaliBERT model performed best and provided an accuracy of 86% with a 0.82&#x202F;F1 score. <xref ref-type="bibr" rid="ref2">Ahmed et al. (2023)</xref> also did Bengali sentiment analysis. E-commerce sentiment classification is done by using transformer-based and transfer learning models. A total of three models are applied for the analysis: LSTM, GRU, and BengaliBERT. For binary classification and multiclass classification, the highest accuracy was 94.5 and 88.78% in BengaliBERT, respectively. After that, <xref ref-type="bibr" rid="ref54">Tareq et al. (2023)</xref> worked with cross-linguistic contextual understanding on Bengali-English code-mixed sentiment analysis. Several machine learning and deep learning models were applied with word embedding models to analyze the data. Among them, XGBoost with the code-mixed Fasttext model gained the best F1 score of 0.87. <xref ref-type="bibr" rid="ref22">Haque et al. (2023)</xref> researched Bengali social media comments for multiclass sentiment classification by using machine learning models. 42,036 Facebook comments trained with features like TF-IDF, CV, and Word2Sequence are applied to several machine learning and deep learning models. Among all the models, CLSTM with the Word2Sequence model performed better than all with an accuracy of 87.80%.</p>
<p>In the part of Explainable AI (XAI), there are few studies done. If we see <xref ref-type="bibr" rid="ref33">Kawakura et al. (2022)</xref> utilized Explainable AI (XAI) techniques based on SHAP, LIME, and LightGBM to analyze agricultural worker datasets. These systems use sensors that are attached to worker&#x2019;s bodies to gather information about how they move in farming. Data scientists use Python programs on devices to look at farm movements and find patterns that can help train farmers. After that, <xref ref-type="bibr" rid="ref20">Dieber and Kirrane (2020)</xref> also worked with Explainable AI to investigate the use of LIME. This study evaluates different mathematical procedures to examine how LIME can be used to understand decisions in fields like healthcare and self-driving cars, testing the comprehensibility of LIME&#x2019;s results. We looked over XAI papers that we used in our work, and LIME is a similar method that we used in our study.</p>
</sec>
<sec id="sec3">
<label>3</label>
<title>Data and methodology</title>
<p>This section provides details about our dataset and the models used in the research. In section 3, we disclosed several contribution details through subsections. In 3.1, mention of the Dataset description. 3.2 describes how the data was collected. Subsection 3.3 is the steps of data pre-processing. In 3.5, ML models are summarized. Finally, 3.6 LSTM describes its layers.</p>
<p><xref ref-type="fig" rid="fig1">Figure 1</xref> illustrates the workflow of our study, including how we collected data and applied feature engineering. It also highlights the number of classes and types of data present in the dataset. After selecting the necessary features, we utilized ML classifiers, LSTM, and Transformer models, along with Explainable AI techniques. Additionally, the research incorporated the use of N-grams and TF-IDF for feature extraction and analysis.</p>
<fig position="float" id="fig1">
<label>Figure 1</label>
<caption><p>Architecture of the methodology for Bengali Title Classification.</p></caption>
<graphic xlink:href="frai-08-1537432-g001.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Diagram illustrating a machine learning model for classification. It includes feature engineering from certain and uncertain data, followed by supervised, sequential, and transformer learning. Data is tokenized and converted to vectors before classification. Categories are national, international, and sports. The process incorporates TF-IDF term frequency, explainable AI, and LIME for model interpretation. A table shows dataset titles in Bengali with categories. An input-output flowchart represents data processing steps, and a robot icon denotes AI assistance.</alt-text>
</graphic>
</fig>
<sec id="sec4">
<label>3.1</label>
<title>Dataset collection</title>
<p>We prepared a raw dataset by ourselves. Bengali news articles are published on the newspaper websites of various newspapers as e-papers (<xref ref-type="bibr" rid="ref55">Timeline, n.d.</xref>). The dataset is prepared manually from these four newspapers: Prothom Alo, Ittefaq, Jugantor, and Kaler Konto. We have uploaded the dataset on Mendeley, and it is publicly available (<xref ref-type="bibr" rid="ref32">Julkar Naeen and Sourav Kumar Das, 2024</xref>).</p>
<p>We ensured proper data annotation by accurately labeling the text data with clear and consistent guidelines. This approach minimized ambiguity and maintained the integrity of the dataset, facilitating effective training and evaluation of the text classification model. Regular quality checks were performed to verify annotation accuracy, ensuring reliability. This meticulous process enhanced the model&#x2019;s ability to learn and deliver precise predictions.</p>
</sec>
<sec id="sec5">
<label>3.2</label>
<title>Dataset description</title>
<p>The dataset is about news articles from newspapers. The dataset has four attributes: title, publisher, newspaper name, and publication date. A total of 6,150 titles are taken in the dataset. <xref ref-type="fig" rid="fig2">Figure 2</xref> provides visualization of the dataset&#x2019;s quantity and distribution based on the three classes. 2089 titles are classified as national, 2008 are classified as sports, and 2053 as international.</p>
<fig position="float" id="fig2">
<label>Figure 2</label>
<caption><p>Class distribution and quantity visualization of three classes.</p></caption>
<graphic xlink:href="frai-08-1537432-g002.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">A donut chart titled "Class Distribution" with three segments: National (34% - 2089), Sports (32.7% - 2008), and International (33.4% - 2053). Each segment is represented by a different color.</alt-text>
</graphic>
</fig>
<p>National: Political, social, economic, and cultural news pertaining to events in Bangladesh.</p>
<p>International: News about events and developments taking place outside Bangladesh, but typically of regional or international significance.</p>
<p>Sports: Updates on local and global sporting competitions, teams, and players&#x2019; performances.</p>
<p><xref ref-type="table" rid="tab1">Table 1</xref> shows a sample of the dataset where the title column has the titles and the category column has the category of the titles. An additional column shows the translated English of titles.</p>
<table-wrap position="float" id="tab1">
<label>Table 1</label>
<caption><p>Sample of the total dataset of Bengali news titles.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Title</th>
<th align="center" valign="top">Category</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">&#x0987;&#x09A8;&#x09CD;&#x09A6;&#x09CB;&#x09A8;&#x09C7;&#x09B6;&#x09BF;&#x09AF;&#x09BC;&#x09BE; &#x09B8;&#x09AB;&#x09B0;&#x09C7; &#x09AF;&#x09BE;&#x099A;&#x09CD;&#x099B;&#x09C7;&#x09A8; &#x0986;&#x09A8;&#x09CB;&#x09AF;&#x09BC;&#x09BE;&#x09B0; &#x0987;&#x09AC;&#x09CD;&#x09B0;&#x09BE;&#x09B9;&#x09BF;&#x09AE;<break/>Anwar Ibrahim is visiting Indonesia</td>
<td align="center" valign="top">International</td>
</tr>
<tr>
<td align="left" valign="top">&#x09AA;&#x09C7;&#x09B2;&#x09C7;&#x09B0; &#x09B8;&#x09C7;&#x0987; &#x09AA;&#x09C7;&#x09A8;&#x09BE;&#x09B2;&#x09CD;&#x099F;&#x09BF; &#x09A0;&#x09C7;&#x0995;&#x09BE;&#x09A8;&#x09CB; &#x09AC;&#x09B2; &#x098F;&#x0996;&#x09A8;&#x09CB; &#x09A4;&#x09BE;&#x09B0; &#x09B8;&#x0982;&#x0997;&#x09CD;&#x09B0;&#x09B9;&#x09C7; Pele&#x2019;s penalty save ball is still in his collection</td>
<td align="center" valign="top">Sports</td>
</tr>
<tr>
<td align="left" valign="top">&#x0995;&#x09BE;&#x09B2;&#x09BF;&#x09AF;&#x09BC;&#x09BE;&#x0995;&#x09C8;&#x09B0;&#x09C7; &#x0985;&#x09B8;&#x09CD;&#x09A4;&#x09CD;&#x09B0; &#x09A0;&#x09C7;&#x0995;&#x09BF;&#x09AF;&#x09BC;&#x09C7; &#x09AC;&#x09CD;&#x09AF;&#x09AC;&#x09B8;&#x09BE;&#x09AF;&#x09BC;&#x09C0;&#x09B0; &#x09AC;&#x09BE;&#x09A1;&#x09BC;&#x09BF;&#x09A4;&#x09C7; &#x09A1;&#x09BE;&#x0995;&#x09BE;&#x09A4;&#x09BF; Armed robbery at businessman&#x2019;s house in Kaliakore</td>
<td align="center" valign="top">National</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>We aimed to gather approximately 2,000 samples for each class since we know that transfer learning models are data-hungry and would perform well with additional data. While reviewing related research, we noted that some research works made use of much less data. Trying to do better than these works, we consciously tried to gather more data than these benchmarks. After collecting some 2,000 samples per class&#x2014;6,150 data points in all&#x2014;we could see a definite improvement in model performance. Even though 6,150 samples are not particularly large, it was enough for our experimental aims. Also, the data was hand-labeled, so the process of collection was laborious and difficult.</p>
</sec>
<sec id="sec6">
<label>3.3</label>
<title>Dataset pre-processing</title>
<p>Processing the data is the crucial part for cleaning the data and making ready for training ML and DL models. <xref ref-type="fig" rid="fig3">Figure 3</xref> presents the steps followed in data pre-processing. This pre-processing step includes data cleaning by removing unnecessary items from the dataset, and then removing stopwords, tokenizers, stemming, null value handling, removing duplicate values, small texts that have no meaning and punctuations, and non-Bengali characters, then removed stopwords, then used Lancaster stemming and tokenized the dataset. For example, &#x201C;&#x09AD;&#x09BE;&#x09B7;&#x09BE;&#x09B6;&#x09BF;&#x0995;&#x09CD;&#x09B7;&#x0995; &#x09AE;&#x09BF;&#x09A5;&#x09BF;&#x09B2;&#x09BE;&#x201D;(Language teacher Mithila). This sentence will not help models identify their category. Steps of pre-processing the data the following order.</p>
<list list-type="simple">
<list-item><p>a. Convert Data Types: First of all, the total dataset is converted to a string type so that models can learn easily. Provides consistency for text models (e.g., BERT, LSTM) that require string inputs.</p></list-item>
<list-item><p>b. Remove Duplicate Row: In case there were any duplicate values, these steps removed all the duplicate values or data if there existed any in the dataset. Removes redundancy, avoiding model bias to overrepresented samples.</p></list-item>
<list-item><p>c. Remove Small Text: A title that has a length of less than four words seems not to make sense or is not understandable. For this situation, a small text of titles that consists of fewer than four words was removed. Brief messages tend to miss contextual information, damaging model.</p></list-item>
<list-item><p>d. Remove Punctuation, Link, Emoji (No Character): Punctuation marks, links, and non-character items create problems for the machine in learning. So non-characters and punctuation are removed. It reduces tokenization noise, particularly for subword models like BERT. Stripping punctuation is particularly necessary for agglutinative languages like Bengali, where suffixes are meaningful but extraneous symbols are not.</p></list-item>
</list>
<p>Some of the punctuations are &#x2018;,&#x2019;, &#x2018;!&#x2019;, &#x2018;?&#x2019;, &#x2018;&#x0964;&#x2019;.</p>
<list list-type="simple">
<list-item><p>e. Remove Non-Bengali Character: Since it&#x2019;s a Bengali dataset and the total research is on Bengali text, having non-Bengali characters in the data means an anomaly. So, all the non-Bengali characters are removed. It trains the model&#x2019;s capacity on Bengali language patterns, free from interference due to mixed-language noise. Essential for monolingual tasks like sentiment analysis or topic classification.</p></list-item>
</list>
<fig position="float" id="fig3">
<label>Figure 3</label>
<caption><p>Preprocessing of Bengal news titles using the necessary steps.</p></caption>
<graphic xlink:href="frai-08-1537432-g003.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Flowchart illustrating text preprocessing steps: Convert Datatype, Remove Duplicate Value, Remove Small Text, Remove Punctuation and Non-Character, Tokenizing, Lancaster Stemming, Remove Stopwords, and Remove Non-Bengali Character. Arrows indicate direction and process order.</alt-text>
</graphic>
</fig>
<p>Remove Stopwords: Stopwords are usually those words that do not have significant meaning in Bengali (<xref ref-type="bibr" rid="ref38">Luhn, 1958</xref>), so these cannot be taken as a tokenizer. Thus, these are noisy data. Moreover, stopwords ain&#x2019;t verbs or tenses and so generate ambiguity. <xref ref-type="table" rid="tab2">Table 2</xref> shows some of the Bengali stopwords. While processing Bengali text, these stopwords are removed for models to understand Bengali text and perform better. This step reduces noise. These noises are created because of several uses of these words. <xref ref-type="table" rid="tab3">Table 3</xref> has three columns showing the results of removing the stopwords from the sentences.</p>
<list list-type="simple">
<list-item><p>f. Stemming: In NLP, stemming means bringing words to their root form. Using this technique, words are reduced to their base form. In this research, Bengali text is produced for training models so that machines can understand it more accurately. <xref ref-type="table" rid="tab4">Table 4</xref> shows the sentences before and after stemming. An additional column was added to the previous tables. This research is conducted using Lancaster stemming (<xref ref-type="bibr" rid="ref42">Paice, 1990</xref>) on the dataset. In Lancaster stemming, there is no change in the sentences, which are shown in <xref ref-type="table" rid="tab4">Table 4</xref>. Stemming covers morphological variation in Bengali (i.e., verb conjugations, plural markers), grouping semantically similar words together. Improves model efficiency but over-stems in some cases.</p></list-item>
<list-item><p>g. Tokenizer: Tokenizing refers to splitting the sentences into raw units such as words (<xref ref-type="bibr" rid="ref50">Salton, 1983</xref>). It helps to transform unprocessed text data into a more structured. Allows out-of-vocabulary words via Byte-Pair Encoding (BPE), crucial for compound words in Bengali.</p></list-item>
</list>
<table-wrap position="float" id="tab2">
<label>Table 2</label>
<caption><p>Some stop words in Bengali.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Bengali stopwords</th>
<th align="center" valign="top">Verbatim</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">&#x098F;&#x0987;</td>
<td align="center" valign="top">This</td>
</tr>
<tr>
<td align="left" valign="top">&#x098F;&#x09AC;&#x0982;</td>
<td align="center" valign="top">And</td>
</tr>
<tr>
<td align="left" valign="top">&#x098F;&#x0995;&#x099F;&#x09BF;</td>
<td align="center" valign="top">A/An</td>
</tr>
<tr>
<td align="left" valign="top">&#x09A4;&#x09CB;</td>
<td align="center" valign="top">That</td>
</tr>
<tr>
<td align="left" valign="top">&#x09A4;&#x09BE;&#x09B9;&#x09B2;&#x09C7;</td>
<td align="center" valign="top">Then</td>
</tr>
<tr>
<td align="left" valign="top">&#x0995;&#x09BF;&#x09A8;&#x09CD;&#x09A4;&#x09C1;</td>
<td align="center" valign="top">But</td>
</tr>
<tr>
<td align="left" valign="top">&#x0995;&#x09AC;&#x09C7;</td>
<td align="center" valign="top">When</td>
</tr>
<tr>
<td align="left" valign="top">&#x09AF;&#x09BE;</td>
<td align="center" valign="top">Which</td>
</tr>
<tr>
<td align="left" valign="top">&#x0995;&#x09CB;&#x09A8;&#x09CB;</td>
<td align="center" valign="top">Any</td>
</tr>
<tr>
<td align="left" valign="top">&#x0995;&#x09BF;&#x099B;&#x09C1;</td>
<td align="center" valign="top">Some</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap position="float" id="tab3">
<label>Table 3</label>
<caption><p>Samples of stopwords removed from the dataset.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Title</th>
<th align="left" valign="top">After removing stopwords</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">&#x098F;&#x0995; &#x09AC;&#x09BE;&#x09B0; &#x099C;&#x09C1;&#x09A4;&#x09CB;&#x09B0; &#x09AB;&#x09BF;&#x09A4;&#x09C7; &#x09AC;&#x09C7;&#x0981;&#x09A7;&#x09C7;&#x0987; &#x09E7; &#x0995;&#x09CB;&#x099F;&#x09BF; &#x09E8;&#x09E9; &#x09B2;&#x09BE;&#x0996;<break/>1 crore 23 lakhs for tying shoelaces once</td>
<td align="center" valign="top">&#x098F;&#x0995; &#x099C;&#x09C1;&#x09A4;&#x09CB;&#x09B0; &#x09AB;&#x09BF;&#x09A4;&#x09C7; &#x09AC;&#x09C7;&#x0981;&#x09A7;&#x09C7;&#x0987; &#x09E7; &#x09E8;&#x09E9; &#x09B2;&#x09BE;&#x0996;</td>
</tr>
<tr>
<td align="left" valign="top">&#x09A4;&#x09C7;&#x09B2; &#x0989;&#x09A4;&#x09CD;&#x09A4;&#x09CB;&#x09B2;&#x09A8;&#x09C7; &#x099A;&#x09C0;&#x09A8;&#x09BE; &#x09AA;&#x09CD;&#x09B0;&#x09A4;&#x09BF;&#x09B7;&#x09CD;&#x09A0;&#x09BE;&#x09A8;&#x09C7;&#x09B0; &#x09B8;&#x0999;&#x09CD;&#x0997;&#x09C7; &#x09A4;&#x09BE;&#x09B2;&#x09C7;&#x09AC;&#x09BE;&#x09A8;&#x09C7;&#x09B0; &#x099A;&#x09C1;&#x0995;&#x09CD;&#x09A4;&#x09BF; Taliban deal with Chinese companies to extract oil</td>
<td align="center" valign="top">&#x09A4;&#x09C7;&#x09B2; &#x0989;&#x09A4;&#x09CD;&#x09A4;&#x09CB;&#x09B2;&#x09A8;&#x09C7; &#x099A;&#x09C0;&#x09A8;&#x09BE; &#x09AA;&#x09CD;&#x09B0;&#x09A4;&#x09BF;&#x09B7;&#x09CD;&#x09A0;&#x09BE;&#x09A8;&#x09C7;&#x09B0; &#x09A4;&#x09BE;&#x09B2;&#x09C7;&#x09AC;&#x09BE;&#x09A8;&#x09C7;&#x09B0; &#x099A;&#x09C1;&#x0995;&#x09CD;&#x09A4;&#x09BF;</td>
</tr>
<tr>
<td align="left" valign="top">&#x09A8;&#x09A4;&#x09C1;&#x09A8; &#x09B8;&#x09CD;&#x09AC;&#x09AA;&#x09CD;&#x09A8; &#x09A8;&#x09BF;&#x09AF;&#x09BC;&#x09C7; &#x09E8;&#x09E6;&#x09E8;&#x09E9; &#x09B8;&#x09BE;&#x09B2;&#x0995;&#x09C7; &#x09AC;&#x09B0;&#x09A3; &#x09AC;&#x09BF;&#x09B6;&#x09CD;&#x09AC;&#x09AC;&#x09BE;&#x09B8;&#x09C0;&#x09B0;<break/>People of the world welcome the year 2023 with new dreams</td>
<td align="center" valign="top">&#x09B8;&#x09CD;&#x09AC;&#x09AA;&#x09CD;&#x09A8; &#x09A8;&#x09BF;&#x09AF;&#x09BC;&#x09C7; &#x09E8;&#x09E6;&#x09E8;&#x09E9; &#x09B8;&#x09BE;&#x09B2;&#x0995;&#x09C7; &#x09AC;&#x09B0;&#x09A3; &#x09AC;&#x09BF;&#x09B6;&#x09CD;&#x09AC;&#x09AC;&#x09BE;&#x09B8;&#x09C0;&#x09B0;</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap position="float" id="tab4">
<label>Table 4</label>
<caption><p>Sample of the titles before and after stemming.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Title</th>
<th align="left" valign="top">After lancaster stemming</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">&#x09E8;&#x09E7; &#x09AC;&#x099B;&#x09B0; &#x09AD;&#x09BE;&#x09B0;&#x09A4; &#x099B;&#x09C7;&#x09B2;&#x09C7;&#x0995;&#x09C7; &#x09AB;&#x09C7;&#x09B0;&#x09A4; &#x09AA;&#x09C7;&#x09B2;&#x09C7;&#x09A8; &#x09AE;&#x09BE; &#x09AC;&#x09BE;&#x09AC;&#x09BE;<break/>21-year-old Indian parents got their son back</td>
<td align="center" valign="top">&#x09E8;&#x09E7; &#x09AC;&#x099B;&#x09B0; &#x09AD;&#x09BE;&#x09B0;&#x09A4; &#x099B;&#x09C7;&#x09B2;&#x09C7;&#x0995;&#x09C7; &#x09AB;&#x09C7;&#x09B0;&#x09A4; &#x09AA;&#x09C7;&#x09B2;&#x09C7;&#x09A8; &#x09AE;&#x09BE; &#x09AC;&#x09BE;&#x09AC;&#x09BE;</td>
</tr>
<tr>
<td align="left" valign="top">&#x09B8;&#x09CD;&#x09AC;&#x09BE;&#x09A7;&#x09C0;&#x09A8;&#x09A4;&#x09BE; &#x09A6;&#x09BF;&#x09AC;&#x09B8;&#x09C7;&#x09B0; &#x09AA;&#x09CD;&#x09B0;&#x09BE;&#x0995;&#x09CD;&#x0995;&#x09BE;&#x09B2;&#x09C7; &#x09AF;&#x09C1;&#x0995;&#x09CD;&#x09A4;&#x09B0;&#x09BE;&#x09B7;&#x09CD;&#x099F;&#x09CD;&#x09B0;&#x09C7;&#x09B0; &#x09AC;&#x09A8;&#x09CD;&#x09A6;&#x09C1;&#x0995; &#x09B9;&#x09BE;&#x09AE;&#x09B2;&#x09BE;&#x09AF;&#x09BC; &#x09A8;&#x09BF;&#x09B9;&#x09A4; &#x09E7;&#x09EB;<break/>15 killed in gun attack in the US on the eve of<break/>Independence Day</td>
<td align="center" valign="top">&#x09B8;&#x09CD;&#x09AC;&#x09BE;&#x09A7;&#x09C0;&#x09A8;&#x09A4;&#x09BE; &#x09A6;&#x09BF;&#x09AC;&#x09B8;&#x09C7;&#x09B0; &#x09AA;&#x09CD;&#x09B0;&#x09BE;&#x0995;&#x09CD;&#x0995;&#x09BE;&#x09B2;&#x09C7; &#x09AF;&#x09C1;&#x0995;&#x09CD;&#x09A4;&#x09B0;&#x09BE;&#x09B7;&#x09CD;&#x099F;&#x09CD;&#x09B0;&#x09C7;&#x09B0; &#x09AC;&#x09A8;&#x09CD;&#x09A6;&#x09C1;&#x0995; &#x09B9;&#x09BE;&#x09AE;&#x09B2;&#x09BE;&#x09AF;&#x09BC; &#x09A8;&#x09BF;&#x09B9;&#x09A4; &#x09E7;&#x09EB;</td>
</tr>
<tr>
<td align="left" valign="top">&#x09AE;&#x09BE;&#x0995;&#x09C7; &#x09AC;&#x09BE;&#x0981;&#x099A;&#x09BE;&#x09A4;&#x09C7; &#x09AD;&#x09BE;&#x0987;&#x09AF;&#x09BC;&#x09C7;&#x09B0; &#x09B9;&#x09BE;&#x09A4;&#x09C7; &#x09AC;&#x09CB;&#x09A8; &#x0996;&#x09C1;&#x09A8;<break/>Sister killed by brother to save mother</td>
<td align="center" valign="top">&#x09AE;&#x09BE;&#x0995;&#x09C7; &#x09AC;&#x09BE;&#x0981;&#x099A;&#x09BE;&#x09A4;&#x09C7; &#x09AD;&#x09BE;&#x0987;&#x09AF;&#x09BC;&#x09C7;&#x09B0; &#x09B9;&#x09BE;&#x09A4;&#x09C7; &#x09AC;&#x09CB;&#x09A8; &#x0996;&#x09C1;&#x09A8;</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>Sample Sentence: &#x201C;&#x0986;&#x0987;&#x09A8;&#x099C;&#x09C0;&#x09AC;&#x09C0; &#x09B9;&#x09A4;&#x09CD;&#x09AF;&#x09BE; &#x09AE;&#x09BE;&#x09AE;&#x09B2;&#x09BE;&#x09AF;&#x09BC; &#x0987;&#x09AE;&#x09B0;&#x09BE;&#x09A8;&#x0996;&#x09BE;&#x09A8;&#x0995;&#x09C7; &#x09B8;&#x09C1;&#x09AA;&#x09CD;&#x09B0;&#x09BF;&#x09AE; &#x0995;&#x09CB;&#x09B0;&#x09CD;&#x099F;&#x09C7; &#x09A4;&#x09B2;&#x09AC;.&#x201D;</p>
<p>Interpreted Sample: &#x201C;Imran Khan summoned to Supreme Court in lawyer murder case.&#x201D;</p>
<p>Token Words: [&#x201C;&#x0986;&#x0987;&#x09A8;&#x099C;&#x09C0;&#x09AC;&#x09C0; (lawyer),&#x201D; &#x201C;&#x09B9;&#x09A4;&#x09CD;&#x09AF;&#x09BE; (murder),&#x201D; &#x201C;&#x09AE;&#x09BE;&#x09AE;&#x09B2;&#x09BE;&#x09AF;&#x09BC; (in case),&#x201D; &#x201C;&#x0987;&#x09AE;&#x09B0;&#x09BE;&#x09A8;&#x0996;&#x09BE;&#x09A8;&#x0995;&#x09C7; (Imran Khan),&#x201D; &#x201C;&#x09B8;&#x09C1;&#x09AA;&#x09CD;&#x09B0;&#x09BF;&#x09AE; (Supreme),&#x201D; &#x201C;&#x0995;&#x09CB;&#x09B0;&#x09CD;&#x099F;&#x09C7; (Court),&#x201D; &#x201C;&#x09A4;&#x09B2;&#x09AC; (summoned)&#x201D;].</p>
</sec>
<sec id="sec7">
<label>3.4</label>
<title>Data summary</title>
<p>The summary produces the length of the words and the length of the characters. Also, the total number of sentences for different classes, the total number of words for each class, the total number of unique words for each class, and their number. These details help us to select the appropriate models for the research, as well as the NLP techniques.</p>
<p><xref ref-type="fig" rid="fig4">Figures 4</xref>, <xref ref-type="fig" rid="fig5">5</xref>. show the length-frequency distribution of words and characters, respectively.</p>
<fig position="float" id="fig4">
<label>Figure 4</label>
<caption><p>Length of the word frequency distribution.</p></caption>
<graphic xlink:href="frai-08-1537432-g004.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Bar chart titled "Length-Frequency Distribution" showing the frequency of words by length. Peaks at word lengths four and six with frequencies over 1200. Frequencies decrease at lengths two, eight, and ten.</alt-text>
</graphic>
</fig>
<fig position="float" id="fig5">
<label>Figure 5</label>
<caption><p>Length of the character frequency distribution.</p></caption>
<graphic xlink:href="frai-08-1537432-g005.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Bar chart titled "Length-Frequency Distribution" showing frequency of character lengths from 0 to 50. Frequencies range from 0 to 200, peaking around length 30. Bars are blue.</alt-text>
</graphic>
</fig>
<p>Moreover, <xref ref-type="fig" rid="fig6">Figure 6</xref> shows that 2056 sentences are classified as national, 2015 sentences are in the international category, and 1989 sentences are sports-related in the dataset after the preprocessing.</p>
<fig position="float" id="fig6">
<label>Figure 6</label>
<caption><p>Data statistics after processing the dataset.</p></caption>
<graphic xlink:href="frai-08-1537432-g006.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">A heatmap titled "Data Statistics" displays data for three categories: National, International, and Sports. It shows total sentences, total words, and unique words for each category. The National category has 2,056 sentences, 12,964 words, and 5,751 unique words. International has 2,015 sentences, 11,838 words, and 4,765 unique words. Sports has 1,989 sentences, 8,379 words, and 3,788 unique words. The color gradient from light to dark blue represents increasing numbers.</alt-text>
</graphic>
</fig>
<p>The total number of national words is 12,964, and 5,751 of the words are unique. 8,379 words are in the sports category, and 3,788 of the words are unique. In the international category, there are 11,838 words, and 4,765 words are unique. The sports have fewer unique words among the three classes. On the other hand, the national category had the most unique words (<xref ref-type="fig" rid="fig7">Figure 7</xref>). There is a simple bar chart of the data statistics for each class.</p>
<fig position="float" id="fig7">
<label>Figure 7</label>
<caption><p>Data statistic bar chart for each class.</p></caption>
<graphic xlink:href="frai-08-1537432-g007.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Bar chart titled "Data Statistics" displaying values for three categories: National, International, and Sports. Each category shows values for Total Sentences, Total Words, and Unique Words. Total Words have the highest values across all categories, followed by Unique Words, with Total Sentences having the lowest.</alt-text>
</graphic>
</fig>
<p><xref ref-type="table" rid="tab5">Table 5</xref> shows the n-gram distributions of the dataset. The frequency of each word in the text is converted into a vector with the help of the TF-IDF. The Ingram range can specify the size of the n-grams. Value 1, 1 shows the unigram where ngrams have a single word, Bigram produces 2 words, and Trigram produces 3 words, which are shown in <xref ref-type="table" rid="tab5">Table 5</xref>.</p>
<table-wrap position="float" id="tab5">
<label>Table 5</label>
<caption><p>N-grams for the news titles based on TF-IDF.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Title</th>
<th align="center" valign="top">Unigram</th>
<th align="center" valign="top">Bigram</th>
<th align="center" valign="top">Trigram</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">&#x09AD;&#x09BE;&#x09B0;&#x09A4; &#x09AE;&#x09B9;&#x09BE;&#x09B8;&#x09BE;&#x0997;&#x09B0;&#x09C7; &#x099C;&#x09B2;&#x09A6;&#x09B8;&#x09CD;&#x09AF;&#x09C1;&#x09A6;&#x09C7;&#x09B0; &#x0995;&#x09AC;&#x09B2;&#x09C7; &#x09AC;&#x09BE;&#x0982;&#x09B2;&#x09BE;&#x09A6;&#x09C7;&#x09B6;&#x09BF; &#x099C;&#x09BE;&#x09B9;&#x09BE;&#x099C;</td>
<td align="center" valign="top">(&#x2018;&#x09B0;&#x09A4;&#x2019;, 1)</td>
<td align="center" valign="top">(&#x2018;&#x09B0;&#x09A4; &#x09AE;&#x09B9;&#x2019;, 1)</td>
<td align="center" valign="top">(&#x2018;&#x09B0;&#x09A4; &#x09AE;&#x09B9; &#x0997;&#x09B0;&#x2019;, 1)</td>
</tr>
<tr>
<td align="left" valign="top">Bangladeshi ship captured by pirates in</td>
<td align="center" valign="top">(&#x2018;&#x09AE;&#x09B9;&#x2019;, 1)</td>
<td align="center" valign="top">(&#x2018;&#x09AE;&#x09B9; &#x0997;&#x09B0;&#x2019;, 1)</td>
<td align="center" valign="top">(&#x2018;&#x09AE;&#x09B9; &#x0997;&#x09B0; &#x099C;&#x09B2;&#x09A6;&#x09B8;&#x2019;, 1)</td>
</tr>
<tr>
<td align="left" valign="top">Indian Ocean</td>
<td align="center" valign="top">(&#x2018;&#x0997;&#x09B0;&#x2019;, 1)</td>
<td align="center" valign="top">(&#x2018;&#x0997;&#x09B0; &#x099C;&#x09B2;&#x09A6;&#x09B8;&#x2019;, 1)</td>
<td align="center" valign="top">(&#x2018;&#x0997;&#x09B0; &#x099C;&#x09B2;&#x09A6;&#x09B8; &#x0995;&#x09AC;&#x09B2;&#x2019;, 1)</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="sec8">
<label>3.5</label>
<title>Machine learning</title>
<p>ML algorithms such as Logistic regression, Decision tree, K-nearest neighbors, Random-forest, Support vector machine, Multi. Naive Bayes and Stochastic gradient descent are applied for experimental results on the dataset.</p>
<sec id="sec9">
<label>3.5.1</label>
<title>Logistic regression</title>
<p>In supervised learning, Logistic Regression is a statistical way to predict the output based on the given input variables (<xref ref-type="bibr" rid="ref13">Cramer, 2002</xref>). By using a logistic function, the output maps the values to 0 or 1. The algorithm assumes the output is a linear combination of input variables.</p>
</sec>
<sec id="sec10">
<label>3.5.2</label>
<title>Decision tree</title>
<p>To make a decision, this tree-based algorithm works recursively, whereby the input space is divided into different groups according to the values of features used by each node to construct tree-like decisions for specific attributes (<xref ref-type="bibr" rid="ref7">Belson, 1959</xref>). One disadvantage of this algorithm is that it overfits noisy data and hence requires pruning strategies.</p>
</sec>
<sec id="sec11">
<label>3.5.3</label>
<title>Random forest</title>
<p>It is a powerful ensemble learning technique that uses a variety of decision trees during the training process. Every tree is trained on some part of the data, and some features are randomly selected. For the classification, the prediction is made by combining outputs from individual trees (<xref ref-type="bibr" rid="ref5">Alzubi et al., 2020</xref>). The usage of random forests is extensive because they are powerful and can handle large data sets with numerous dimensions (<xref ref-type="bibr" rid="ref11">Breiman, 2001</xref>).</p>
</sec>
<sec id="sec12">
<label>3.5.4</label>
<title>Multi. Naive Bayes</title>
<p>Multinomial Naive Bayes is a version of the Naive Bayes algorithm especially created for text classification problems in which features are individual words, frequencies, or any other discrete features (<xref ref-type="bibr" rid="ref46">Rennie, 2001</xref>). In this algorithm, features are independent when it comes to class values. The simplicity of Multinomial Naive Bayes does not make it less effective, as it has been able to work efficiently against many different types of problems. It performs quickly and handles large feature sets too.</p>
</sec>
<sec id="sec13">
<label>3.5.5</label>
<title>K-Nearest neighbors</title>
<p>KNN is an uncomplicated machine-learning algorithm. In the feature space, it ultimately decides the dot class or value by looking at how classes or values are assigned to other data dots that are next to it. It is very simple to understand and can be done easily, but only if the K-neighbours parameter is picked correctly, which states how many neighboring points (K) from each testing sample should be used in making predictions about other classes based on their attributes, along with the kind of measure among them (<xref ref-type="bibr" rid="ref14">Cunningham and Delany, 2021</xref>).</p>
</sec>
<sec id="sec14">
<label>3.5.6</label>
<title>Support vector machine</title>
<p>The SVM algorithm is useful for carrying out classification tasks. It can establish the ideal hyperplane for separating varied classes by making the margin between support vectors as wide as possible (<xref ref-type="bibr" rid="ref10">Boswell, 2002</xref>). SVM employs kernel tricks to address non-linearity between classes and features in very many dimensions. For detecting spams and sentiment analysis, this model gives the best output (<xref ref-type="bibr" rid="ref44">Qiqieh et al., 2025</xref>).</p>
</sec>
</sec>
<sec id="sec15">
<label>3.6</label>
<title>Long short-term memory (LSTM)</title>
<p>The LSTM model used in this study is sequential (<xref ref-type="bibr" rid="ref26">Hochreiter, 1997</xref>). It processes sequential data into a sequence, and a more abstract representation and gives an output suitable for classification or regression. The embed dim is 64. The input dim is 5,000. The dropout is 0.2, and the recurrent dropout is 0.4. The total parameters in the embedding layer is 320,000. The spatial dropout layer uses 0 input units as a dropout for the embedding layer&#x2019;s input, which occurs during each training time. The spatial dropout 1D is 0.4. The LSTM is a type of RNN model that is quite popular in Bengali NLP. In this layer, the output shape has been changed into a 64-dimensional vector. There are a total of 33,024 parameters in the LSTM layer. The dense layer is completely linked, and it is activated using softmax. The output shape is (None,3). The total number of parameters in the dense layer is 195. Activation Function (SOFTMAX) converts a real number into a probability distribution. The total number of parameters in the model is 353,219.</p>
</sec>
<sec id="sec16">
<label>3.7</label>
<title>Transformer model for classification</title>
<p>Many earlier studies in natural language processing have successfully used transfer learning models like XLM-RoBERTa base, Multilingual BERT, and DistilBERT. We studied these earlier studies and their accuracy on different tasks in detail and resolved to use these models (<xref ref-type="bibr" rid="ref27">Hoque et al., 2024</xref> and <xref ref-type="bibr" rid="ref48">Roy et al., 2023</xref>).</p>
<sec id="sec17">
<label>3.7.1</label>
<title>Multilingual BERT</title>
<p>The Bert-based multilingual case is a pre-trained model of the BERT, developed for handling various languages (<xref ref-type="bibr" rid="ref34">Kenton and Toutanova, 2019</xref>). Masked language modeling is the main goal of this model. This allows the model to understand the languages in the deep contextual meaning.</p>
<p>The base architecture of this model is the 12 transformer layers. Each hidden layer contains 768 neurons. With 12 attention heads in each encoder layer, the self-attention layer focuses on multiple parts of the input. The feed-forward network of each encoder block has 3,072 intermediate sizes. The dropout of attention and the hidden layer are the same, 0.1. The vocabulary size is 119,547. The model can process bidirectional inputs.</p>
<p>This model can fine-tune language tasks by adapting multilingual embeddings to language. The downstream NLP tasks such as sentiment analysis, classification, or NER.</p>
</sec>
<sec id="sec18">
<label>3.7.2</label>
<title>DistilBERT</title>
<p>DistilBERT is a small and faster version of the BERT model on text classification (<xref ref-type="bibr" rid="ref51">Sanh, 2019</xref>). It takes tokenized text as input and processes it through the DistilBERT encoder, pooling layer, and dense layer. This model has 4 layers of architecture. First is the Input layer, which takes tokenized text as input with a shape of (batch_size, sequence_length). Then, the input text is passed to the encoder layer and processed to generate contextual meaning. Then the pooling layer applies mean pooling to make a sequence-level representation. After that, the dense layer maps the pool and is classified using softmax or sigmoid.</p>
</sec>
<sec id="sec19">
<label>3.7.3</label>
<title>XLM-RoBERTa-base</title>
<p>The XLM-RoBERTa-base model is a transformer-based language model designed for natural language processing (NLP) tasks (<xref ref-type="bibr" rid="ref12">Conneau, 2019</xref>). It is a lightweight version of the XLM-RoBERTa model, optimized for efficiency while maintaining strong performance. Pre-trained on data in 100 languages, this model is highly versatile for multilingual tasks. Fine-tuning was conducted with a learning rate of 2e-5 times, and the given learning rate was chosen to ensure stability during optimization. The training process spanned five epochs, with a batch size of 16 for both training and evaluation, and input sequences tokenized to a maximum length of 512 tokens. The classification task involved three labels.</p>
<p>The model architecture includes several noteworthy configurations. It employs a dropout probability of 0.1 in the attention mechanism to reduce overfitting, and the GELU activation function is used in the feedforward layers. The hidden layer size is 768, with an intermediate feedforward size of 3,072. The model consists of 12 transformer layers, each with 12 attention heads. A hidden layer dropout probability of 0.1 was also applied. The model supports up to 514 tokens and has a vocabulary size of 250,002, with a total of 278,045,955 trainable parameters. During training, Weights &#x0026; Biases (W&#x0026;B) logging was disabled.</p>
<p>The training dataset contained 4,920 examples, processed with a gradient accumulation step of 1, leading to a total of 1,540 optimization steps provided in <xref ref-type="disp-formula" rid="EQ1">Equation 1</xref>.</p>
<disp-formula id="EQ1"><mml:math id="M1"><mml:mi mathvariant="italic">Optimization</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="italic">Steps</mml:mi><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>N</mml:mi><mml:mi>u</mml:mi><mml:mi>m</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="italic">Examples</mml:mi><mml:mo>&#x00D7;</mml:mo><mml:mi>N</mml:mi><mml:mi>u</mml:mi><mml:mi>m</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="italic">Epochs</mml:mi></mml:mrow><mml:mrow><mml:mi mathvariant="italic">Total</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="italic">Train</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="italic">Batch</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="italic">Size</mml:mi></mml:mrow></mml:mfrac><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mn>4920</mml:mn><mml:mo>&#x00D7;</mml:mo><mml:mn>5</mml:mn><mml:mspace width="0.25em"/></mml:mrow><mml:mn>16</mml:mn></mml:mfrac><mml:mo>=</mml:mo><mml:mn>1540</mml:mn></mml:math><label>(1)</label></disp-formula>
<p>The evaluation was conducted on a dataset comprising 1,230 examples, also with a batch size of 16. This phase took approximately 33 s, processing 37 samples provided in <xref ref-type="disp-formula" rid="E1">Equation 2</xref>.</p>
<disp-formula id="E1"><mml:math id="M2"><mml:mi mathvariant="normal">Samples</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="normal">per</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="normal">second</mml:mi><mml:mo>=</mml:mo><mml:mi mathvariant="normal">Steps</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="normal">per</mml:mi><mml:mspace width="0.25em"/><mml:mi mathvariant="normal">second</mml:mi><mml:mo>&#x00D7;</mml:mo><mml:mi mathvariant="normal">Batch size</mml:mi><mml:mo>=</mml:mo><mml:mn>2.321</mml:mn><mml:mo>&#x00D7;</mml:mo><mml:mn>16</mml:mn><mml:mo>&#x2248;</mml:mo><mml:mn>37.1</mml:mn></mml:math><label>(2)</label></disp-formula>
<p>The fact that the minimum losses that were recorded in training and the validation stages further indicates that the model has successfully avoided overfitting and that it can easily generalize to the unknown data. As a result, the XLMRoBERTa-base architecture proved to be highly successful and reliable in all the measures that were considered. The construction of the input representation of each of the tokens is defined in <xref ref-type="disp-formula" rid="E2">Equation 3</xref>, where token representations are augmented with positional representations to produce contextualized input representations. Attention mechanism, contextual inter-token relationships, is expressed by means of <xref ref-type="disp-formula" rid="E3">Equation 4</xref>. The nonlinear transformation of the attention outputs, which is done by the feed-forward network, is defined in <xref ref-type="disp-formula" rid="E4">Equation 5</xref>. Introducing layer normalization and residual connections make gradient propagation stable, which is supported in <xref ref-type="disp-formula" rid="E5">Equation 6</xref>. Lastly, the masked language modeling goal objective used in pretraining is elaborated in <xref ref-type="disp-formula" rid="E6">Equation 7</xref>.</p>
<p>Input Representation:</p>
<disp-formula id="E2"><mml:math id="M3"><mml:msub><mml:mi>x</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mi>E</mml:mi><mml:mfenced open="(" close=")"><mml:msub><mml:mi>t</mml:mi><mml:mi>i</mml:mi></mml:msub></mml:mfenced><mml:mo>+</mml:mo><mml:mi>P</mml:mi><mml:mfenced open="(" close=")"><mml:mi>i</mml:mi></mml:mfenced></mml:math><label>(3)</label></disp-formula>
<p>Self-Attention Mechanism:</p>
<disp-formula id="E3"><mml:math id="M4"><mml:msup><mml:mi>Z</mml:mi><mml:mfenced open="(" close=")"><mml:mi>i</mml:mi></mml:mfenced></mml:msup><mml:mo>=</mml:mo><mml:mi mathvariant="italic">softmax</mml:mi><mml:mo stretchy="false">(</mml:mo><mml:mfrac><mml:mrow><mml:msup><mml:mi>Q</mml:mi><mml:mfenced open="(" close=")"><mml:mi>l</mml:mi></mml:mfenced></mml:msup><mml:msup><mml:mi>K</mml:mi><mml:msup><mml:mfenced open="(" close=")"><mml:mi>i</mml:mi></mml:mfenced><mml:mi>t</mml:mi></mml:msup></mml:msup></mml:mrow><mml:msqrt><mml:msub><mml:mi>d</mml:mi><mml:mi>k</mml:mi></mml:msub></mml:msqrt></mml:mfrac></mml:math><label>(4)</label></disp-formula>
<p>Feed-Forward Network (FFN):</p>
<disp-formula id="E4"><mml:math id="M5"><mml:mi>F</mml:mi><mml:mi>F</mml:mi><mml:msub><mml:mi>N</mml:mi><mml:mfenced open="(" close=")"><mml:mi>z</mml:mi></mml:mfenced></mml:msub><mml:mo>=</mml:mo><mml:mi mathvariant="italic">ReLU</mml:mi><mml:mfenced open="(" close=")"><mml:mrow><mml:mi>z</mml:mi><mml:msub><mml:mi>W</mml:mi><mml:mn>1</mml:mn></mml:msub><mml:mo>+</mml:mo><mml:msub><mml:mi>b</mml:mi><mml:mn>1</mml:mn></mml:msub></mml:mrow></mml:mfenced><mml:msub><mml:mi>W</mml:mi><mml:mn>2</mml:mn></mml:msub><mml:mo>+</mml:mo><mml:msub><mml:mi>b</mml:mi><mml:mn>2</mml:mn></mml:msub></mml:math><label>(5)</label></disp-formula>
<p>Layer Normalization and Residual Connection:</p>
<disp-formula id="E5"><mml:math id="M6"><mml:mi>X</mml:mi><mml:mi>l</mml:mi><mml:mfenced open="(" close=")"><mml:mrow><mml:mi>l</mml:mi><mml:mo>+</mml:mo><mml:mn>1</mml:mn></mml:mrow></mml:mfenced><mml:mo>=</mml:mo><mml:mi mathvariant="italic">LayerNorm</mml:mi><mml:mfenced open="(" close=")"><mml:mrow><mml:mi>X</mml:mi><mml:mfenced open="(" close=")"><mml:mi>l</mml:mi></mml:mfenced><mml:mo>+</mml:mo><mml:mi>Z</mml:mi><mml:mfenced open="(" close=")"><mml:mi>l</mml:mi></mml:mfenced></mml:mrow></mml:mfenced></mml:math><label>(6)</label></disp-formula>
<p>Masked Language Model Objective:</p>
<disp-formula id="E6"><mml:math id="M7"><mml:msub><mml:mi>L</mml:mi><mml:mrow><mml:mi>M</mml:mi><mml:mi>L</mml:mi><mml:mi>M</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:munder><mml:mo stretchy="true">&#x2211;</mml:mo><mml:mrow><mml:mi>i</mml:mi><mml:mo>&#x2208;</mml:mo><mml:mi>M</mml:mi></mml:mrow></mml:munder></mml:mstyle><mml:mi mathvariant="italic">logP</mml:mi><mml:mfenced open="(" close=")"><mml:mrow><mml:msub><mml:mi>t</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo stretchy="true">|</mml:mo><mml:mi>X</mml:mi><mml:mo>&#x2212;</mml:mo><mml:mi>M</mml:mi></mml:mrow></mml:mfenced></mml:math><label>(7)</label></disp-formula>
</sec>
</sec>
</sec>
<sec sec-type="results" id="sec20">
<label>4</label>
<title>Results and discussion</title>
<p>Results and discussion are the key part of the research, which provides a complete perspective on the findings of the research. This part of the paper presents a detailed analysis that has been provided from the dataset, different models&#x2019; performance, and evaluation of their scores.</p>
<sec id="sec21">
<label>4.1</label>
<title>Model evaluation</title>
<p>Based on accuracy, precision, recall, and F1-score, different models are evaluated and compared in their performance. Accuracy, precision, recall, and F1-score are calculated with the formulas given below.</p>
<p>Accuracy: In <xref ref-type="disp-formula" rid="E7">Equation 8</xref>, the accuracy of a model refers to the proportion of predictions it makes, calculated as the ratio of positives and true negatives to all positive and negative observations.</p>
<disp-formula id="E7"><mml:math id="M8"><mml:mi mathvariant="italic">Accuracy</mml:mi><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>+</mml:mo><mml:mi>T</mml:mi><mml:mi>N</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>+</mml:mo><mml:mi>T</mml:mi><mml:mi>N</mml:mi><mml:mo>+</mml:mo><mml:mi>F</mml:mi><mml:mi>P</mml:mi><mml:mo>+</mml:mo><mml:mi>F</mml:mi><mml:mi>N</mml:mi></mml:mrow></mml:mfrac></mml:math><label>(8)</label></disp-formula>
<p>Precision: As expressed in <xref ref-type="disp-formula" rid="E8">Equation 9</xref>, it is made up of the ratio of the number of correct model predictions, which focuses on the precision in relation to positive predictions.</p>
<disp-formula id="E8"><mml:math id="M9"><mml:mi mathvariant="italic">Precision</mml:mi><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>+</mml:mo><mml:mi>F</mml:mi><mml:mi>P</mml:mi></mml:mrow></mml:mfrac></mml:math><label>(9)</label></disp-formula>
<p>Recall: From <xref ref-type="disp-formula" rid="E9">Equation 10</xref>, Recall shows how much a given model can identify all instances of a given class correctly, and it is usually calculated as the number of true positives divided by the sum of all numbers that are represented by true positives and false negatives.</p>
<disp-formula id="E9"><mml:math id="M10"><mml:mi mathvariant="italic">Recall</mml:mi><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi></mml:mrow><mml:mrow><mml:mi>T</mml:mi><mml:mi>P</mml:mi><mml:mo>+</mml:mo><mml:mi>F</mml:mi><mml:mi>N</mml:mi></mml:mrow></mml:mfrac></mml:math><label>(10)</label></disp-formula>
<p>F1 Score: Based on <xref ref-type="disp-formula" rid="E10">Equation 11</xref>, it measure statisticians use when analyzing data. When we combine the two measures (precision and recall), it is called a hybrid measure. The F1-score is calculated as the average of precision and recall.</p>
<disp-formula id="E10"><mml:math id="M11"><mml:mi>F</mml:mi><mml:mn>1</mml:mn><mml:mo>=</mml:mo><mml:mfrac><mml:mrow><mml:mn>2</mml:mn><mml:mo>&#x2217;</mml:mo><mml:mi mathvariant="italic">Precision</mml:mi><mml:mo>&#x2217;</mml:mo><mml:mi mathvariant="italic">Recall</mml:mi></mml:mrow><mml:mrow><mml:mi mathvariant="italic">Precision</mml:mi><mml:mo>+</mml:mo><mml:mi mathvariant="italic">Recall</mml:mi></mml:mrow></mml:mfrac></mml:math><label>(11)</label></disp-formula>
</sec>
<sec id="sec22">
<label>4.2</label>
<title>Comparison of models&#x2019; performance</title>
<p><xref ref-type="table" rid="tab6">Table 6</xref> shows the comparison of performance by ML models with LSTM. Among the ML models, Multi. Naive Bayes did a good performance with 85.22% accuracy. The best model among the DL and ML models is LSTM, which is the top performer. SVM also performed well, almost matching Naive Bayes. Logistic Regression also delivered a strong performance. However, the other models exhibited some issues. The precision was higher than the accuracy and other metrics. <xref ref-type="table" rid="tab7">Table 7</xref> shows the performance of BERT models, XLM-Roberta base outperformed with an accuracy of 91.38%. Here, the best two performing models are the XLMRoberta Base and Multilingual BERT, with an accuracy of 91.38 and 87.64%. But the DistilBERT&#x2019;s performance was poor, infect very poor. Reason is that DistilBERT is trained on English data (Wikipedia + BookCorpus) only. Since, it&#x2019;s not pre-trained on Bengali data or multilingual data, tokenization mismatched for Bengali text. DistilBERT is not familiar with Bengali vocabulary, this led to not understanding the tokens, thus giving this poor performance.</p>
<table-wrap position="float" id="tab6">
<label>Table 6</label>
<caption><p>Comparison of ML models&#x2019; performance.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Models</th>
<th align="center" valign="top">Accuracy (%)</th>
<th align="center" valign="top">Precision (%)</th>
<th align="center" valign="top">Recall (%)</th>
<th align="center" valign="top">F1-Score (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">LSTM</td>
<td align="center" valign="top">86.25</td>
<td align="center" valign="top">86.25</td>
<td align="center" valign="top">86.25</td>
<td align="center" valign="top">86.25</td>
</tr>
<tr>
<td align="left" valign="top">Multi. Naive Bayes</td>
<td align="center" valign="top">85.22</td>
<td align="center" valign="top">85.78</td>
<td align="center" valign="top">85.22</td>
<td align="center" valign="top">85.16</td>
</tr>
<tr>
<td align="left" valign="top">SVM</td>
<td align="center" valign="top">84.71</td>
<td align="center" valign="top">84.98</td>
<td align="center" valign="top">84.71</td>
<td align="center" valign="top">84.76</td>
</tr>
<tr>
<td align="left" valign="top">Logistic Regression</td>
<td align="center" valign="top">81.36</td>
<td align="center" valign="top">82.26</td>
<td align="center" valign="top">81.36</td>
<td align="center" valign="top">81.43</td>
</tr>
<tr>
<td align="left" valign="top">KNN</td>
<td align="center" valign="top">76.55</td>
<td align="center" valign="top">77.22</td>
<td align="center" valign="top">76.55</td>
<td align="center" valign="top">76.43</td>
</tr>
<tr>
<td align="left" valign="top">SGD</td>
<td align="center" valign="top">74.00</td>
<td align="center" valign="top">79.29</td>
<td align="center" valign="top">74.00</td>
<td align="center" valign="top">73.79</td>
</tr>
<tr>
<td align="left" valign="top">Random Forest</td>
<td align="center" valign="top">73.28</td>
<td align="center" valign="top">79.52</td>
<td align="center" valign="top">73.28</td>
<td align="center" valign="top">73.40</td>
</tr>
<tr>
<td align="left" valign="top">Decision Tree</td>
<td align="center" valign="top">71.56</td>
<td align="center" valign="top">74.24</td>
<td align="center" valign="top">71.56</td>
<td align="center" valign="top">71.76</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap position="float" id="tab7">
<label>Table 7</label>
<caption><p>Comparison of deep learning and transformer models performance.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Models</th>
<th align="center" valign="top">Accuracy (%)</th>
<th align="center" valign="top">Precision (%)</th>
<th align="center" valign="top">Recall (%)</th>
<th align="center" valign="top">F1-Score (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">XLM-Roberta_Base</td>
<td align="center" valign="top">91.38</td>
<td align="center" valign="top">91.38</td>
<td align="center" valign="top">91.41</td>
<td align="center" valign="top">91.39</td>
</tr>
<tr>
<td align="left" valign="top">Multilingual BERT</td>
<td align="center" valign="top">87.64</td>
<td align="center" valign="top">87.79</td>
<td align="center" valign="top">87.68</td>
<td align="center" valign="top">87.71</td>
</tr>
<tr>
<td align="left" valign="top">DistilBERT</td>
<td align="center" valign="top">53.0</td>
<td align="center" valign="top">52.0</td>
<td align="center" valign="top">53.0</td>
<td align="center" valign="top">48.9</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>From <xref ref-type="table" rid="tab6">Tables 6</xref>, <xref ref-type="table" rid="tab7">7</xref>, the best-performing model was XLM-Roberta Base, also the 2nd best model is also a transformer-based model. Compared to the ML and DL models, the transformer-based model was quite better, also showing prominent performance.</p>
</sec>
<sec id="sec23">
<label>4.3</label>
<title>Confusion matrix comaprison</title>
<p>A confusion matrix is an effective way to evaluate classifier models. From a confusion matrix, a clear vision can be achieved of the outcome of the model and whether the acquired accuracy is valid. Issues such as underfitting or overfitting can also be identified and addressed using a confusion matrix.</p>
<p><xref ref-type="table" rid="tab8">Table 8</xref> illustrates the comparative performance of four top-performing models analyzed in this study: XLM-RoBERTa, Multilingual BERT, LSTM, and Multinomial Naive Bayes, respectively. Among these, Multinomial Naive Bayes and LSTM ranked fourth and third, respectively, while Multilingual BERT emerged as the second-best model. The highest-performing model was XLM-RoBERTa, achieving an accuracy of 91.38%. Notably, the BERT-based models demonstrated superior performance compared to both machine learning (ML) and deep learning (DL) models, as evidenced by higher true positive (TP) rates and lower false negative (FN) and false positive (FP) rates. Within the BERT-based models, XLM-RoBERTa outperformed Multilingual BERT, further validating its superior performance in terms of TP and FP metrics.</p>
<table-wrap position="float" id="tab8">
<label>Table 8</label>
<caption><p>Confusion matrix for best-performing ML, DL (LSTM), and Transformer models.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Model name</th>
<th align="center" valign="top">Class</th>
<th align="center" valign="top">TP</th>
<th align="center" valign="top">TN</th>
<th align="center" valign="top">FP</th>
<th align="center" valign="top">FN</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle" rowspan="3">XLM-RoBERTa BASE</td>
<td align="center" valign="top">0</td>
<td align="center" valign="top">334</td>
<td align="center" valign="top">711</td>
<td align="center" valign="top">54</td>
<td align="center" valign="top">65</td>
</tr>
<tr>
<td align="center" valign="top">1</td>
<td align="center" valign="top">353</td>
<td align="center" valign="top">682</td>
<td align="center" valign="top">65</td>
<td align="center" valign="top">64</td>
</tr>
<tr>
<td align="center" valign="top">2</td>
<td align="center" valign="top">317</td>
<td align="center" valign="top">775</td>
<td align="center" valign="top">41</td>
<td align="center" valign="top">31</td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Multilingual BERT</td>
<td align="center" valign="top">0</td>
<td align="center" valign="top">359</td>
<td align="center" valign="top">757</td>
<td align="center" valign="top">66</td>
<td align="center" valign="top">48</td>
</tr>
<tr>
<td align="center" valign="top">1</td>
<td align="center" valign="top">365</td>
<td align="center" valign="top">749</td>
<td align="center" valign="top">56</td>
<td align="center" valign="top">60</td>
</tr>
<tr>
<td align="center" valign="top">2</td>
<td align="center" valign="top">354</td>
<td align="center" valign="top">802</td>
<td align="center" valign="top">30</td>
<td align="center" valign="top">44</td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">LSTM</td>
<td align="center" valign="top">0</td>
<td align="center" valign="top">334</td>
<td align="center" valign="top">711</td>
<td align="center" valign="top">54</td>
<td align="center" valign="top">65</td>
</tr>
<tr>
<td align="center" valign="top">1</td>
<td align="center" valign="top">353</td>
<td align="center" valign="top">682</td>
<td align="center" valign="top">65</td>
<td align="center" valign="top">64</td>
</tr>
<tr>
<td align="center" valign="top">2</td>
<td align="center" valign="top">317</td>
<td align="center" valign="top">775</td>
<td align="center" valign="top">41</td>
<td align="center" valign="top">31</td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Multinominal Na&#x00EF;ve Bayes</td>
<td align="center" valign="top">0</td>
<td align="center" valign="top">336</td>
<td align="center" valign="top">709</td>
<td align="center" valign="top">62</td>
<td align="center" valign="top">57</td>
</tr>
<tr>
<td align="center" valign="top">1</td>
<td align="center" valign="top">316</td>
<td align="center" valign="top">729</td>
<td align="center" valign="top">29</td>
<td align="center" valign="top">90</td>
</tr>
<tr>
<td align="center" valign="top">2</td>
<td align="center" valign="top">340</td>
<td align="center" valign="top">718</td>
<td align="center" valign="top">81</td>
<td align="center" valign="top">25</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>Managing Polysemy and Syntactic Ambiguity: Our BERT-based model tackles polysemous words and syntactic ambiguity in Bengali news headlines using contextual embeddings and self-attention mechanisms. In contrast to static representations, BERT disambiguates words dynamically (e.g., &#x201C;&#x09AA;&#x09A6;&#x201D; as &#x201C;position&#x201D; or &#x201C;foot&#x201D;) based on bidirectional context analysis. In the case of syntactic ambiguity (e.g., free word order in &#x201C;&#x09A6;&#x09C0;&#x09B0;&#x09CD;&#x0998;&#x09A6;&#x09BF;&#x09A8;&#x09C7;&#x09B0; &#x09AC;&#x09C8;&#x09B7;&#x09AE;&#x09CD;&#x09AF;&#x09C7;&#x09B0; &#x0995;&#x09BE;&#x09B0;&#x09A3;&#x09C7;&#x0987; &#x09B8;&#x09B9;&#x09BF;&#x0982;&#x09B8;&#x201D;), multi-head attention settles dependencies based on weighting pertinent token relationships.</p>
</sec>
<sec id="sec24">
<label>4.4</label>
<title>Title classification explanation of XAI-based LIME</title>
<sec id="sec25">
<label>4.4.1</label>
<title>Local interpretable model-agnostic explanations (LIME)</title>
<p>We select the ground data coordinates, input them into the black box scheme, and observe the corresponding outputs. This technique assesses the new data based on its proximity to the original coordinate points. Consequently, it fits an alternative model, such as linear regression, to the modified sample set using the derived weights. Henceforth, any original data point can be interpreted using the newly developed explanatory model.</p>
<p>Explainability in a model refers to the capacity to comprehend and interpret the processes by which the model generates its predictions or decisions. While a model may demonstrate high performance and accuracy across various tasks, it often functions as a &#x201C;black box,&#x201D; making it challenging to ascertain the rationale behind specific predictions or outcomes. LIME (<xref ref-type="bibr" rid="ref47">Ribeiro et al., 2016</xref>) initiates the process by altering the Bengali newspaper, introducing subtle modifications such as rearranging, excluding, or inserting words. This approach aims to evaluate the model&#x2019;s sensitivity to deviations in the input data. The altered Bengali newspaper is then input into the transformer-based black-box model, which produces a prediction. LIME subsequently identifies the key features from the altered instances that are most affected by these modifications. These local surrogate models approximate the behavior of the black-box model in the proximity of the selected instance, providing insight into the underlying decision-making process. The knowledge derived from these local models is then leveraged to explain the predictions made by the black-box model on the original dataset. Typically, these explanations emphasize the terms or critical features that have a significant impact on the model&#x2019;s decision-making, thereby enhancing the interpretability of the predictive mechanism. Through a systematic and iterative methodology, LIME aids in uncovering the complexities of machine learning models and promotes their transparency, thereby fostering greater confidence in the accuracy of their predictions.</p>
<p>To be more specific, an explanation for a data point x is a model g that minimizes the locality-aware loss L(f, g, &#x03C0;x) associated with how well &#x2018;g&#x2019; approximates the original function f in its neighborhood &#x03C0;x while maintaining low complexity, denoted by the model provided in <xref ref-type="disp-formula" rid="E11">Equation 12</xref>.</p>
<disp-formula id="E11"><mml:math id="M12"><mml:mi mathvariant="italic">argmi</mml:mi><mml:msub><mml:mi>n</mml:mi><mml:mi>g</mml:mi></mml:msub><mml:mi>L</mml:mi><mml:mfenced open="(" close=")" separators=",,"><mml:mi>f</mml:mi><mml:mi>g</mml:mi><mml:msub><mml:mi>&#x03C0;</mml:mi><mml:mi>x</mml:mi></mml:msub></mml:mfenced><mml:mo>+</mml:mo><mml:mi mathvariant="normal">&#x03A9;</mml:mi><mml:mfenced open="(" close=")"><mml:mi>g</mml:mi></mml:mfenced></mml:math><label>(12)</label></disp-formula>
<p><xref ref-type="table" rid="tab9">Table 9</xref> reveals &#x2018;&#x0989;&#x09A4;&#x2019; (weight: &#x2212;0.0742) as the most significant negative influencer, followed by &#x2018;&#x099C;&#x09A8;&#x2019; (+0.0640) as the strongest positive contributor. Moderate influences include &#x2018;&#x098F;&#x09AC;&#x2019; (&#x2212;0.0624) and &#x2018;&#x09AC;&#x09B6;&#x2019; (+0.0349), while &#x2018;&#x09AF;&#x09CE;&#x2019; and &#x2018;&#x0987;&#x0993;&#x09AF;&#x2019; show weaker positive effects. These weights indicate each term&#x2019;s relative importance in the model&#x2019;s decision-making process, with absolute values quantifying their impact magnitude. These forms of Bengali words are called &#x201C;&#x09A7;&#x09BE;&#x09A4;&#x09C1;,&#x201D; which means verbal root. These words do not carry specific meaning; affixes make them meaningful.</p>
<table-wrap position="float" id="tab9">
<label>Table 9</label>
<caption><p>Weight quantifying their impact magnitude relative importance in the model&#x2019;s decision-making.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Bengali feature</th>
<th align="center" valign="top">Raw weight</th>
<th align="center" valign="top">Absolute weight</th>
<th align="center" valign="top">Interpretation</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">&#x0989;&#x09A4;</td>
<td align="center" valign="top">&#x2212;0.0742</td>
<td align="center" valign="top">0.0742</td>
<td align="center" valign="top">Strong Negative</td>
</tr>
<tr>
<td align="left" valign="top">&#x099C;&#x09A8;</td>
<td align="center" valign="top">+0.0640</td>
<td align="center" valign="top">0.0640</td>
<td align="center" valign="top">Strong Positive</td>
</tr>
<tr>
<td align="left" valign="top">&#x098F;&#x09AC;</td>
<td align="center" valign="top">&#x2212;0.0624</td>
<td align="center" valign="top">0.0624</td>
<td align="center" valign="top">Moderate Negative</td>
</tr>
<tr>
<td align="left" valign="top">&#x09AC;&#x09B6;</td>
<td align="center" valign="top">+0.0349</td>
<td align="center" valign="top">0.0349</td>
<td align="center" valign="top">Moderate Positive</td>
</tr>
<tr>
<td align="left" valign="top">&#x09AF;&#x09CE;</td>
<td align="center" valign="top">+0.0277</td>
<td align="center" valign="top">0.0277</td>
<td align="center" valign="top">Weak Positive</td>
</tr>
<tr>
<td align="left" valign="top">&#x0987;&#x0993;&#x09AF;</td>
<td align="center" valign="top">+0.0089</td>
<td align="center" valign="top">0.0089</td>
<td align="center" valign="top">Minimal Positive</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>In <xref ref-type="fig" rid="fig8">Figures 8a</xref>,<xref ref-type="fig" rid="fig8">b</xref>, two examples are given for comparing how the model assessed misclassified titles. <xref ref-type="fig" rid="fig8">Figure 8a</xref> shows the sentence that has been classified correctly. It says &#x201C;&#x099A;&#x09C0;&#x09A8; &#x09A4;&#x09BE;&#x0987;&#x0993;&#x09AF;&#x09BC;&#x09BE;&#x09A8; &#x0989;&#x09A4;&#x09CD;&#x09A4;&#x09C7;&#x099C;&#x09A8;&#x09BE; &#x098F;&#x09AC;&#x0982; &#x09AC;&#x09BF;&#x09B6;&#x09CD;&#x09AC;&#x09B6;&#x09BE;&#x09A8;&#x09CD;&#x09A4;&#x09BF;&#x09B0; &#x09AD;&#x09AC;&#x09BF;&#x09B7;&#x09CD;&#x09AF;&#x09CE;&#x201D; means &#x201C;China-Taiwan Tensions and the Future of World Peace&#x201D; is an international predicted as international as well. The probability score and the weighted value are presented in the figure, which is 0.78, and the weighted values are shown in <xref ref-type="table" rid="tab9">Table 9</xref>. <xref ref-type="fig" rid="fig8">Figure 8b</xref> shows the misclassified sentence and how it is assessed. &#x201C;&#x099B;&#x09BE;&#x09A6;&#x0996;&#x09CB;&#x09B2;&#x09BE; &#x09AC;&#x09BE;&#x09B8;&#x09C7;&#x09B0; &#x0995;&#x09A5;&#x09BE; &#x09AE;&#x09A8;&#x09C7; &#x0995;&#x09B0;&#x09BF;&#x09AF;&#x09BC;&#x09C7; &#x09A6;&#x09BF;&#x09AF;&#x09BC;&#x09C7;&#x099B;&#x09C7;&#x09A8; &#x09B8;&#x09BE;&#x09A8;&#x099C;&#x09BF;&#x09A6;&#x09BE;&#x201D; means &#x201C;Sanjida reminded me of an open-top bus&#x201D; misclassified as National instead of Sports. Reason is shown in <xref ref-type="fig" rid="fig8">Figure 8b</xref>, that is weight value of each verbal root is positive for national, and the probability for National is 0.59, where the probability for sports is 0.30.</p>
<fig position="float" id="fig8">
<label>Figure 8</label>
<caption><p>Example XAI-based LIME output for XLM-RoBERTa base for two predictions (misclassified vs. classified). <bold>(a)</bold> Correctly classified <bold>(b)</bold> Misclassified.</p></caption>
<graphic xlink:href="frai-08-1537432-g008.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Panel (a) shows prediction probabilities with International at 0.78, National at 0.12, and Sports at 0.10, alongside highlighted Bengali words. Panel (b) displays probabilities with International at 0.11, National at 0.59, and Sports at 0.30, with different highlighted Bengali words.</alt-text>
</graphic>
</fig>
</sec>
</sec>
<sec id="sec26">
<label>4.5</label>
<title>Testing performance with uncertain data</title>
<p>In this subsection, <xref ref-type="table" rid="tab10">Table 10</xref> shows the predicted output of the proposed model. The actual class is the labeled category of the titles, and the prediction column is the predicted class of the model. It seems that the models performed very well on predicting, though there are very few wrong predictions as well. But still, most of the titles are predicted accurately. It is cross-checked for confusing titles, and model mistakes while predicting.</p>
<table-wrap position="float" id="tab10">
<label>Table 10</label>
<caption><p>Prediction performance of the XLM-RoBERTa model on the test data.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Title</th>
<th align="center" valign="top">Actual class</th>
<th align="center" valign="top">Predicted</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">&#x09AA;&#x09BE;&#x09AC;&#x09A8;&#x09BE;&#x09B0; &#x098F;&#x0995; &#x0996;&#x09BE;&#x09AE;&#x09BE;&#x09B0; &#x0989;&#x09AA;&#x09B9;&#x09BE;&#x09B0; &#x098F;&#x09B8;&#x09C7;&#x099B;&#x09C7; &#x09E7;&#x09EF;&#x099F;&#x09BF; &#x099A;&#x09BF;&#x09A4;&#x09CD;&#x09B0;&#x09BE; &#x09B9;&#x09B0;&#x09BF;&#x09A3;<break/>A farm in Pabna was gifted with 19 Chitra deer</td>
<td align="center" valign="top">National</td>
<td align="center" valign="top">National</td>
</tr>
<tr>
<td align="left" valign="top">&#x0987;&#x09B8;&#x09B0;&#x09BE;&#x09AF;&#x09BC;&#x09C7;&#x09B2;&#x09BF; &#x0995;&#x09B0;&#x09CD;&#x09AE;&#x0995;&#x09BE;&#x09A3;&#x09CD;&#x09A1; &#x09AC;&#x09BF;&#x09AC;&#x09C7;&#x099A;&#x09A8;&#x09BE;&#x09AF;&#x09BC; &#x0986;&#x0987;&#x09B8;&#x09BF;&#x099C;&#x09C7;&#x09B0; &#x09B8;&#x09C1;&#x09AA;&#x09BE;&#x09B0;&#x09BF;&#x09B6;<break/>ICJ recommendations regarding Israeli actions</td>
<td align="center" valign="top">International</td>
<td align="center" valign="top">International</td>
</tr>
<tr>
<td align="left" valign="top">&#x09B8;&#x09A4;&#x09CD;&#x09AF; &#x0993;&#x09AF;&#x09BC;&#x09C7;&#x09B7;&#x09CD;&#x099F; &#x0987;&#x09A8;&#x09CD;&#x09A1;&#x09BF;&#x099C;&#x0995;&#x09C7; &#x099B;&#x09BE;&#x09A1;&#x09BC;&#x09BE;&#x0987; &#x09AC;&#x09BF;&#x09B6;&#x09CD;&#x09AC;&#x0995;&#x09BE;&#x09AA;<break/>World Cup without Satya West Indies</td>
<td align="center" valign="top">Sports</td>
<td align="center" valign="top">Sports</td>
</tr>
<tr>
<td align="left" valign="top">&#x09AA;&#x09CD;&#x09B2;&#x09BE;&#x09B8;&#x09CD;&#x099F;&#x09BF;&#x0995;&#x09C7;&#x09B0; &#x09AA;&#x09C1;&#x09A8;&#x09B0;&#x09CD;&#x09AC;&#x09CD;&#x09AF;&#x09AC;&#x09B9;&#x09BE;&#x09B0;&#x09C7; &#x09A8;&#x09BF;&#x09AF;&#x09BC;&#x09A8;&#x09CD;&#x09A4;&#x09CD;&#x09B0;&#x09A3; &#x09A6;&#x09C2;&#x09B7;&#x09A3;<break/>Control pollution in plastic recycling</td>
<td align="center" valign="top">National</td>
<td align="center" valign="top">National</td>
</tr>
<tr>
<td align="left" valign="top">&#x09A6;&#x09C0;&#x09B0;&#x09CD;&#x0998;&#x09A6;&#x09BF;&#x09A8;&#x09C7;&#x09B0; &#x09AC;&#x09C8;&#x09B7;&#x09AE;&#x09CD;&#x09AF;&#x09C7;&#x09B0; &#x0995;&#x09BE;&#x09B0;&#x09A3;&#x09C7;&#x0987; &#x09B8;&#x09B9;&#x09BF;&#x0982;&#x09B8;<break/>Violence is the result of long-standing discrimination</td>
<td align="center" valign="top">International</td>
<td align="center" valign="top">International</td>
</tr>
<tr>
<td align="left" valign="top">&#x09AE;&#x09C2;&#x09B2; &#x0985;&#x09AD;&#x09BF;&#x09AF;&#x09C1;&#x0995;&#x09CD;&#x09A4;&#x09C7;&#x09B0; &#x09AC;&#x09BE;&#x09A1;&#x09BC;&#x09BF;&#x09A4;&#x09C7; &#x09AC;&#x09BF;&#x0995;&#x09CD;&#x09B7;&#x09CB;&#x09AD;&#x0995;&#x09BE;&#x09B0;&#x09C0;&#x09A6;&#x09C7;&#x09B0; &#x0986;&#x0997;&#x09C1;&#x09A8;<break/>Protesters set fire to main accused&#x2019;s house</td>
<td align="center" valign="top">National</td>
<td align="center" valign="top">National</td>
</tr>
<tr>
<td align="left" valign="top">&#x09AA;&#x09C1;&#x09B0;&#x09CB;&#x09A8;&#x09CB; &#x09AC;&#x09A8;&#x09CD;&#x09A7;&#x09C1; &#x0995;&#x09BF;&#x09B8;&#x09BF;&#x099E;&#x09CD;&#x099C;&#x09BE;&#x09B0;&#x0995;&#x09C7; &#x09B8;&#x09CD;&#x09AC;&#x09BE;&#x0997;&#x09A4; &#x099C;&#x09BE;&#x09A8;&#x09BE;&#x09B2;&#x09C7;&#x09A8; &#x099C;&#x09BF;&#x09A8;&#x09AA;&#x09BF;&#x0982;<break/>Xi Jinping welcomes old friend Kissinger</td>
<td align="center" valign="top">International</td>
<td align="center" valign="top">International</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="sec27">
<label>4.6</label>
<title>Comparison of some similar works</title>
<p>The research gap we found that are shown on the Limitation column of the <xref ref-type="table" rid="tab11">Table 11</xref>. It seems that all the works done in Bengali language were conducted on low amount of data, because of that their models did not perform best. We also found the validation of their models on uncertain data were missing or ignored.</p>
<table-wrap position="float" id="tab11">
<label>Table 11</label>
<caption><p>Comparison with some state of work.</p></caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Authors</th>
<th align="center" valign="top">Contribution</th>
<th align="center" valign="top">Methods</th>
<th align="center" valign="top">Accuracy(best)</th>
<th align="center" valign="top">Limitation</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top"><xref ref-type="bibr" rid="ref18">Dhar and Abedin (2021)</xref></td>
<td align="center" valign="top">News title categorization</td>
<td align="center" valign="top">Logistic Regression, KNN, Naive Bayes, Adaboost, SVM</td>
<td align="center" valign="top">73.68% (TF-IDF) 75.78% (Countvect.)</td>
<td align="center" valign="top">Didn&#x2019;t validate model&#x2019;s performance</td>
</tr>
<tr>
<td align="left" valign="top"><xref ref-type="bibr" rid="ref15">Das et al. (2021)</xref></td>
<td align="center" valign="top">Bengali hate speech detection category</td>
<td align="center" valign="top">1Dconvolutional layers, LSTM, GRU-based decoders</td>
<td align="center" valign="top">Attention-based decoder 77%</td>
<td align="center" valign="top">lower accuracy and missing model performance validation</td>
</tr>
<tr>
<td align="left" valign="top"><xref ref-type="bibr" rid="ref36">Khan et al. (2021)</xref></td>
<td align="center" valign="top">Sentiment Analysis</td>
<td align="center" valign="top">SVM KNN, ANN, Random-forest, Naive Bayes</td>
<td align="center" valign="top">62%</td>
<td align="center" valign="top">very poor performance of the models</td>
</tr>
<tr>
<td align="left" valign="top">Our approach</td>
<td align="center" valign="top">Bengali text classification, news titles categorization</td>
<td align="center" valign="top">XLM-Roberta</td>
<td align="center" valign="top">91.38%</td>
<td/>
</tr>
</tbody>
</table>
</table-wrap>
<p><xref ref-type="table" rid="tab11">Table 11</xref> compares some works similar to this study based on their dataset, methods, and findings. This comparison provides a clear idea about the gap between the previous works and this study, and finds a better way to classify Bengali news titles.</p>
<p>The result of this study stands out better compared to the papers of some previous studies listed in <xref ref-type="table" rid="tab11">Table 11</xref>. In this study, the average accuracy was 91.38%. The precision, recall value, and F1-score of this research are the same. These previous papers lacked explanations of model performance, and they did not provide any explanation of their model performance validation.</p>
<p>The results of this study represent an optimized outcome, demonstrating superior performance compared to other works in the field. Many of the reviewed studies reported lower accuracy scores and exhibited minimal differences between accuracy and other performance metrics. In some cases, the results were notably suboptimal, with very low accuracy, potentially due to inadequate preprocessing of the dataset. In this study, the integration of LIME played a pivotal role in evaluating the models&#x2019; understanding of the provided dataset. This approach not only ensured robust model performance but also enhanced interpretability. The outcomes of the proposed model are noteworthy and underline its effectiveness in addressing the research problem.</p>
</sec>
</sec>
<sec id="sec28">
<label>5</label>
<title>Conclusion and future scope</title>
<p>This research focuses on Bengali news article classification employing various machine learning models and deep learning techniques, emphasizing LSTM neural networks and BERT base models. We aimed to improve the classification precision of predefined categories of Bengali news articles to aid future growth in the NLP domain, specifically for low-resource languages such as Bengali. This research emphasizes the effectiveness of deep learning models on text classification tasks, specifically in Bengali. It points out the interpretability of Explainable AI (XAI), ensuring that the methods we apply are not only functional but also well understood by all people. We aim to encourage more research and innovation in the field of NLP, focusing on low-resource languages such as Bengali, by proposing a method of categorizing Bengali news articles and solving its unique issues.</p>
<p>For future works, we will increase our dataset to improve model accuracy, perform fine-tuning and hyperparameter optimization to improve the model&#x2019;s performance, and also integrate additional explainability techniques to further improve model clarity.</p>
<sec id="sec29">
<label>5.1</label>
<title>Experimental setup</title>
<p>Experiments were conducted with a combination of local and cloud computing resources. We worked on a system powered by a 12th Gen Intel(R) Core(TM) i5-12450H processor at 2.00&#x202F;GHz with 16&#x202F;GB of RAM (15.7&#x202F;GB usable), on a 64-bit Windows operating system with an x64-based processor architecture. We utilized Google Colab to train and fine-tune the model on the NVIDIA T4 GPU, which is easily accessible within the Colab environment. Training was completed in a total of around 38&#x202F;min and 55&#x202F;s, with throughput equal to 10.534 training samples/s and 0.659 training steps/s.</p>
</sec>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="sec30">
<title>Data availability statement</title>
<p>The datasets presented in this study can be found in online repositories. The Dataset is publicly available on Mendeley. (doi: <ext-link xlink:href="http://doi.org/10.17632/g6ygmy7s5r.2" ext-link-type="uri">10.17632/g6ygmy7s5r.2</ext-link>).</p>
</sec>
<sec sec-type="author-contributions" id="sec31">
<title>Author contributions</title>
<p>MN: Conceptualization, Data curation, Formal analysis, Methodology, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. SD: Conceptualization, Data curation, Formal analysis, Methodology, Resources, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. SJ: Conceptualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. SK: Supervision, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. NS: Methodology, Writing &#x2013; review &#x0026; editing. MH: Supervision, Writing &#x2013; review &#x0026; editing. Ohidujjaman: Supervision, Writing &#x2013; review &#x0026; editing.</p>
</sec>
<sec sec-type="COI-statement" id="sec33">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="sec34">
<title>Generative AI statement</title>
<p>The author(s) declare that no Gen AI was used in the creation of this manuscript.</p>
<p>Any alternative text (alt text) provided alongside figures in this article has been generated by Frontiers with the support of artificial intelligence and reasonable efforts have been made to ensure accuracy, including review by the authors wherever possible. If you identify any issues, please contact us.</p>
</sec>
<sec sec-type="disclaimer" id="sec35">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="ref1"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Ahmed</surname> <given-names>M. T.</given-names></name> <name><surname>Rahman</surname> <given-names>M.</given-names></name> <name><surname>Nur</surname> <given-names>S.</given-names></name> <name><surname>Islam</surname> <given-names>A. Z. M. T.</given-names></name> <name><surname>Das</surname> <given-names>D.</given-names></name></person-group> (<year>2021</year>). <article-title>Natural language processing and machine learning based cyberbullying detection for Bangla and Romanized Bangla texts</article-title>. <source>TELKOMNIKA (Telecommunication Computing Electronics and Control)</source> <volume>20</volume>, <fpage>89</fpage>&#x2013;<lpage>97</lpage>. doi: <pub-id pub-id-type="doi">10.12928/telkomnika.v20i1.18630</pub-id></mixed-citation></ref>
<ref id="ref2"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Ahmed</surname> <given-names>Z.</given-names></name> <name><surname>Shanto</surname> <given-names>S. S.</given-names></name> <name><surname>Jony</surname> <given-names>A. I.</given-names></name></person-group> (<year>2023</year>). <article-title>Advancement in Bangla sentiment analysis: a comparative study of transformer-based and transfer learning models for ecommerce sentiment classification</article-title>. <source>J. Inf. Syst. Eng. Bus. Intell.</source> <volume>9</volume>, <fpage>181</fpage>&#x2013;<lpage>194</lpage>. doi: <pub-id pub-id-type="doi">10.20473/jisebi.9.2.181-194</pub-id></mixed-citation></ref>
<ref id="ref3"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Al Mahmud</surname> <given-names>T.</given-names></name> <name><surname>Sultana</surname> <given-names>S.</given-names></name> <name><surname>Mondal</surname> <given-names>A.</given-names></name></person-group> (<year>2023</year>). <article-title>A new technique to classification of Bengali news grounded on ML and DL models</article-title>. <source>Int. J. Comput. Appl.</source> <volume>975</volume>:<fpage>8887</fpage>.</mixed-citation></ref>
<ref id="ref4"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Alam</surname> <given-names>T.</given-names></name> <name><surname>Khan</surname> <given-names>A.</given-names></name> <name><surname>Alam</surname> <given-names>F.</given-names></name></person-group> (<year>2020</year>). <article-title>Bangla text classification using transformers</article-title>. <source>arXiv</source>:<fpage>2011.04446</fpage>.</mixed-citation></ref>
<ref id="ref5"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Alzubi</surname> <given-names>O. A.</given-names></name> <name><surname>Alzubi</surname> <given-names>J. A.</given-names></name> <name><surname>Alweshah</surname> <given-names>M.</given-names></name> <name><surname>Qiqieh</surname> <given-names>I.</given-names></name> <name><surname>Al-Shami</surname> <given-names>S.</given-names></name> <name><surname>Ramachandran</surname> <given-names>M.</given-names></name></person-group> (<year>2020</year>). <article-title>An optimal pruning algorithm of classifier ensembles: dynamic programming approach</article-title>. <source>Neural Comput. &#x0026; Applic.</source> <volume>32</volume>, <fpage>16091</fpage>&#x2013;<lpage>16107</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s00521-020-04761-6</pub-id>, PMID: <pub-id pub-id-type="pmid">41020202</pub-id></mixed-citation></ref>
<ref id="ref6"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Aslam</surname> <given-names>N.</given-names></name> <name><surname>Khan</surname> <given-names>I. U.</given-names></name> <name><surname>Mirza</surname> <given-names>S.</given-names></name> <name><surname>AlOwayed</surname> <given-names>A.</given-names></name> <name><surname>Anis</surname> <given-names>F. M.</given-names></name> <name><surname>Aljuaid</surname> <given-names>R. M.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Interpretable machine learning models for malicious domains detection using explainable artificial intelligence (XAI)</article-title>. <source>Sustainability</source> <volume>14</volume>:<fpage>7375</fpage>. doi: <pub-id pub-id-type="doi">10.3390/su14127375</pub-id></mixed-citation></ref>
<ref id="ref7"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Belson</surname> <given-names>W. A.</given-names></name></person-group> (<year>1959</year>). <article-title>Matching and prediction on the principle of biological classification</article-title>. <source>J. R. Stat. Soc.: Ser. C: Appl. Stat.</source> <volume>8</volume>, <fpage>65</fpage>&#x2013;<lpage>75</lpage>. doi: <pub-id pub-id-type="doi">10.2307/2985543</pub-id></mixed-citation></ref>
<ref id="ref8"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Bhowmik</surname> <given-names>N. R.</given-names></name> <name><surname>Arifuzzaman</surname> <given-names>M.</given-names></name> <name><surname>Mondal</surname> <given-names>M. R. H.</given-names></name></person-group> (<year>2022</year>). <article-title>Sentiment analysis on Bangla text using extended lexicon dictionary and deep learning algorithms</article-title>. <source>Array</source> <volume>13</volume>:<fpage>100123</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.array.2021.100123</pub-id>, PMID: <pub-id pub-id-type="pmid">41016916</pub-id></mixed-citation></ref>
<ref id="ref9"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Bhowmik</surname> <given-names>N. R.</given-names></name> <name><surname>Arifuzzaman</surname> <given-names>M.</given-names></name> <name><surname>Mondal</surname> <given-names>M. R. H.</given-names></name> <name><surname>Islam</surname> <given-names>M. S.</given-names></name></person-group> (<year>2021</year>). <article-title>Bangla text sentiment analysis using supervised machine learning with extended lexicon dictionary</article-title>. <source>Nat. Lang. Processing Res.</source> <volume>1</volume>, <fpage>34</fpage>&#x2013;<lpage>45</lpage>. doi: <pub-id pub-id-type="doi">10.2991/nlpr.d.210316.001</pub-id>, PMID: <pub-id pub-id-type="pmid">32175718</pub-id></mixed-citation></ref>
<ref id="ref10"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Boswell</surname> <given-names>D.</given-names></name></person-group> (<year>2002</year>). Introduction to support vector machines. Department of Computer Science and Engineering, University of California San Diego, 11, 16&#x2013;17</mixed-citation></ref>
<ref id="ref11"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Breiman</surname> <given-names>L.</given-names></name></person-group> (<year>2001</year>). <article-title>Random forests</article-title>. <source>Mach. Learn.</source> <volume>45</volume>, <fpage>5</fpage>&#x2013;<lpage>32</lpage>. doi: <pub-id pub-id-type="doi">10.1023/A:1010933404324</pub-id></mixed-citation></ref>
<ref id="ref12"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Conneau</surname> <given-names>A.</given-names></name></person-group> (<year>2019</year>). <article-title>Unsupervised cross-lingual representation learning at scale</article-title>. <source>arXiv</source>:<fpage>1911.02116</fpage>.</mixed-citation></ref>
<ref id="ref13"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Cramer</surname> <given-names>J. S.</given-names></name></person-group> (<year>2002</year>). <source>The origins of logistic regression</source>. <publisher-loc>Hoboken, NJ, USA</publisher-loc>: <publisher-name>Wiley</publisher-name>.</mixed-citation></ref>
<ref id="ref14"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Cunningham</surname> <given-names>P.</given-names></name> <name><surname>Delany</surname> <given-names>S. J.</given-names></name></person-group> (<year>2021</year>). <article-title>K-nearest neighbour classifiers&#x2014;a tutorial</article-title>. <source>ACM Comput. Surv.</source> <volume>54</volume>, <fpage>1</fpage>&#x2013;<lpage>25</lpage>. doi: <pub-id pub-id-type="doi">10.1145/3459665</pub-id></mixed-citation></ref>
<ref id="ref15"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Das</surname> <given-names>A. K.</given-names></name> <name><surname>Al Asif</surname> <given-names>A.</given-names></name> <name><surname>Paul</surname> <given-names>A.</given-names></name> <name><surname>Hossain</surname> <given-names>M. N.</given-names></name></person-group> (<year>2021</year>). <article-title>Bangla hate speech detection on social media using attention-based recurrent neural network</article-title>. <source>J. Intell. Syst.</source> <volume>30</volume>, <fpage>578</fpage>&#x2013;<lpage>591</lpage>. doi: <pub-id pub-id-type="doi">10.1515/jisys-2020-0060</pub-id></mixed-citation></ref>
<ref id="ref16"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Das</surname> <given-names>R. K.</given-names></name> <name><surname>Islam</surname> <given-names>M.</given-names></name> <name><surname>Hasan</surname> <given-names>M. M.</given-names></name> <name><surname>Razia</surname> <given-names>S.</given-names></name> <name><surname>Hassan</surname> <given-names>M.</given-names></name> <name><surname>Khushbu</surname> <given-names>S. A.</given-names></name></person-group> (<year>2023a</year>). <article-title>Sentiment analysis in multilingual context: comparative analysis of machine learning and hybrid deep learning models</article-title>. <source>Heliyon</source> <volume>9</volume>:<fpage>e20281</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.heliyon.2023.e20281</pub-id>, PMID: <pub-id pub-id-type="pmid">37809397</pub-id></mixed-citation></ref>
<ref id="ref17"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Das</surname> <given-names>R. K.</given-names></name> <name><surname>Islam</surname> <given-names>M.</given-names></name> <name><surname>Khushbu</surname> <given-names>S. A.</given-names></name></person-group> (<year>2023b</year>). <article-title>BTSD: a curated transformation of sentence dataset for text classification in Bangla language</article-title>. <source>Data Brief</source> <volume>50</volume>:<fpage>109445</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.dib.2023.109445</pub-id>, PMID: <pub-id pub-id-type="pmid">37577411</pub-id></mixed-citation></ref>
<ref id="ref18"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Dhar</surname> <given-names>P.</given-names></name> <name><surname>Abedin</surname> <given-names>M. Z.</given-names></name></person-group> (<year>2021</year>). <article-title>Bengali news headline categorization using optimized machine learning pipeline</article-title>. <source>Int. J. Inf. Eng. Electron. Bus.</source> <volume>13</volume>, <fpage>15</fpage>&#x2013;<lpage>24</lpage>. doi: <pub-id pub-id-type="doi">10.5815/ijieeb.2021.01.02</pub-id></mixed-citation></ref>
<ref id="ref19"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Dhar</surname> <given-names>B.</given-names></name> <name><surname>Morshed</surname> <given-names>M. N.</given-names></name></person-group> (<year>2022</year>). A comparative analysis of Bangla crime news categorization using most prominent machine learning algorithms (PhD thesis) Sonargaon University (SU).</mixed-citation></ref>
<ref id="ref20"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Dieber</surname> <given-names>J.</given-names></name> <name><surname>Kirrane</surname> <given-names>S.</given-names></name></person-group> (<year>2020</year>). <article-title>Why model why? Assessing the strengths and limitations of LIME</article-title>. <source>arXiv</source>:<fpage>2012.00093</fpage>.</mixed-citation></ref>
<ref id="ref21"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Fouad</surname> <given-names>K. M.</given-names></name> <name><surname>Sabbeh</surname> <given-names>S. F.</given-names></name> <name><surname>Medhat</surname> <given-names>W.</given-names></name></person-group> (<year>2022</year>). <article-title>Arabic fake news detection using deep learning</article-title>. <source>Comput. Mater. Contin.</source> <volume>71</volume>, <fpage>3647</fpage>&#x2013;<lpage>3665</lpage>. doi: <pub-id pub-id-type="doi">10.32604/cmc.2022.021449</pub-id></mixed-citation></ref>
<ref id="ref22"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Haque</surname> <given-names>R.</given-names></name> <name><surname>Islam</surname> <given-names>N.</given-names></name> <name><surname>Tasneem</surname> <given-names>M.</given-names></name> <name><surname>Das</surname> <given-names>A. K.</given-names></name></person-group> (<year>2023</year>). <article-title>Multi-class sentiment classification on Bengali social media comments using machine learning</article-title>. <source>Int. J. Cogn. Comp. Eng.</source> <volume>4</volume>, <fpage>21</fpage>&#x2013;<lpage>35</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.ijcce.2023.01.001</pub-id></mixed-citation></ref>
<ref id="ref23"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hasan</surname> <given-names>M. K.</given-names></name> <name><surname>Islam</surname> <given-names>S. A.</given-names></name> <name><surname>Ejaz</surname> <given-names>M. S.</given-names></name> <name><surname>Alam</surname> <given-names>M. M.</given-names></name> <name><surname>Mahmud</surname> <given-names>N.</given-names></name> <name><surname>Rafin</surname> <given-names>T. A.</given-names></name></person-group> (<year>2023a</year>). <article-title>Classifying Bengali newspaper headlines with advanced deep learning models: LSTM, bi-LSTM, and bi-GRU approaches</article-title>. <source>Asian J. Res. Comput. Sci.</source> <volume>16</volume>, <fpage>372</fpage>&#x2013;<lpage>388</lpage>. doi: <pub-id pub-id-type="doi">10.9734/ajrcos/2023/v16i4398</pub-id>, PMID: <pub-id pub-id-type="pmid">40828949</pub-id></mixed-citation></ref>
<ref id="ref24"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hasan</surname> <given-names>M.</given-names></name> <name><surname>Islam</surname> <given-names>L.</given-names></name> <name><surname>Jahan</surname> <given-names>I.</given-names></name> <name><surname>Meem</surname> <given-names>S. M.</given-names></name> <name><surname>Rahman</surname> <given-names>R. M.</given-names></name></person-group> (<year>2023b</year>). <article-title>Natural language processing and sentiment analysis on Bangla social media comments on Russia&#x2013;Ukraine war using transformers</article-title>. <source>Vietnam J. Comput. Sci.</source> <volume>10</volume>, <fpage>329</fpage>&#x2013;<lpage>356</lpage>. doi: <pub-id pub-id-type="doi">10.1142/S2196888823500021</pub-id></mixed-citation></ref>
<ref id="ref25"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hassan</surname> <given-names>M.</given-names></name> <name><surname>Shakil</surname> <given-names>S.</given-names></name> <name><surname>Moon</surname> <given-names>N. N.</given-names></name> <name><surname>Islam</surname> <given-names>M. M.</given-names></name> <name><surname>Hossain</surname> <given-names>R. A.</given-names></name> <name><surname>Mariam</surname> <given-names>A.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Sentiment analysis on Bangla conversation using machine learning approach</article-title>. <source>Int. J. Elect. Comp. Eng.</source> <volume>12</volume>, <fpage>5562</fpage>&#x2013;<lpage>5572</lpage>. doi: <pub-id pub-id-type="doi">10.11591/ijece.v12i5.pp5562-5572</pub-id></mixed-citation></ref>
<ref id="ref26"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hochreiter</surname> <given-names>S.</given-names></name></person-group> (<year>1997</year>). <article-title>Long short-term memory</article-title>. <source>Neural Comput.</source> doi: <pub-id pub-id-type="doi">10.1162/neco.1997.9.8.1735</pub-id>, PMID: <pub-id pub-id-type="pmid">9377276</pub-id></mixed-citation></ref>
<ref id="ref27"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hoque</surname> <given-names>M. N.</given-names></name> <name><surname>Salma</surname> <given-names>U.</given-names></name> <name><surname>Uddin</surname> <given-names>M. J.</given-names></name> <name><surname>Ahamad</surname> <given-names>M. M.</given-names></name> <name><surname>Aktar</surname> <given-names>S.</given-names></name></person-group> (<year>2024</year>). <article-title>Exploring transformer models in the sentiment analysis task for the under-resource Bengali language</article-title>. <source>Nat. Lang. Processing J.</source> <volume>8</volume>:<fpage>100091</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.nlp.2024.100091</pub-id>, PMID: <pub-id pub-id-type="pmid">41016916</pub-id></mixed-citation></ref>
<ref id="ref28"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Hossain</surname> <given-names>E.</given-names></name> <name><surname>Chaudhary</surname> <given-names>N.</given-names></name> <name><surname>Rifad</surname> <given-names>Z. H.</given-names></name> <name><surname>Hossain</surname> <given-names>B.</given-names></name></person-group> (<year>2020b</year>). Bangla-news-headlines categorization. GitHub.</mixed-citation></ref>
<ref id="ref29"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hossain</surname> <given-names>M. Z.</given-names></name> <name><surname>Rahman</surname> <given-names>M. A.</given-names></name> <name><surname>Islam</surname> <given-names>M. S.</given-names></name> <name><surname>Kar</surname> <given-names>S.</given-names></name></person-group> (<year>2020a</year>). <article-title>Banfakenews: a dataset for detecting fake news in Bangla</article-title>. <source>arXiv</source>:<fpage>2004.08789</fpage>.</mixed-citation></ref>
<ref id="ref30"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hossain</surname> <given-names>M. R.</given-names></name> <name><surname>Sarkar</surname> <given-names>S.</given-names></name> <name><surname>Rahman</surname> <given-names>M.</given-names></name></person-group> (<year>2020c</year>). <article-title>Different machine learning based approaches of baseline and deep learning models for Bengali news categorization</article-title>. <source>Int. J. Comput. Appl.</source> <volume>975</volume>:<fpage>8887</fpage>.</mixed-citation></ref>
<ref id="ref31"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Hussain</surname> <given-names>M. G.</given-names></name> <name><surname>Sultana</surname> <given-names>B.</given-names></name> <name><surname>Rahman</surname> <given-names>M.</given-names></name> <name><surname>Hasan</surname> <given-names>M. R.</given-names></name></person-group> (<year>2023</year>). <article-title>Comparison analysis of Bangla news articles classification using support vector machine and logistic regression</article-title>. <source>TELKOMNIKA (Telecommunication Computing Electronics and Control)</source> <volume>21</volume>, <fpage>584</fpage>&#x2013;<lpage>591</lpage>. doi: <pub-id pub-id-type="doi">10.12928/telkomnika.v21i3.23416</pub-id></mixed-citation></ref>
<ref id="ref32"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Julkar Naeen</surname> <given-names>S. A. J.</given-names></name> <name><surname>Sourav Kumar Das</surname></name></person-group>. (<year>2024</year>). Explainable detection: A transformer-based language modeling approach for Bengali news title classification with comparative explainability analysis using ML &#x0026; DL. Mendeley Data, Version 2.</mixed-citation></ref>
<ref id="ref33"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Kawakura</surname> <given-names>S.</given-names></name> <name><surname>Hirafuji</surname> <given-names>M.</given-names></name> <name><surname>Ninomiya</surname> <given-names>S.</given-names></name> <name><surname>Shibasaki</surname> <given-names>R.</given-names></name></person-group> (<year>2022</year>). <article-title>Analyses of diverse agricultural worker data with explainable artificial intelligence: XAI based on SHAP, LIME, and LightGBM</article-title>. <source>Eur. J. Agric. Food Sci.</source> <volume>4</volume>, <fpage>11</fpage>&#x2013;<lpage>19</lpage>. doi: <pub-id pub-id-type="doi">10.24018/ejfood.2022.4.6.348</pub-id></mixed-citation></ref>
<ref id="ref34"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Kenton</surname> <given-names>J. D. M.-W. C.</given-names></name> <name><surname>Toutanova</surname> <given-names>L. K.</given-names></name></person-group> (<year>2019</year>). &#x201C;BERT: Pre-training of deep bidirectional transformers for language understanding.&#x201D; In <italic>Proceedings of NAACL-HLT (Vol. 1, p. 2). Minneapolis, Minnesota</italic>.</mixed-citation></ref>
<ref id="ref35"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Keya</surname> <given-names>A. J.</given-names></name> <name><surname>Wadud</surname> <given-names>M. A. H.</given-names></name> <name><surname>Mridha</surname> <given-names>M.</given-names></name> <name><surname>Alatiyyah</surname> <given-names>M.</given-names></name> <name><surname>Hamid</surname> <given-names>M. A.</given-names></name></person-group> (<year>2022</year>). <article-title>Augfakebert: handling imbalance through augmentation of fake news using BERT to enhance the performance of fake news classification</article-title>. <source>Appl. Sci.</source> <volume>12</volume>:<fpage>8398</fpage>. doi: <pub-id pub-id-type="doi">10.3390/app12178398</pub-id></mixed-citation></ref>
<ref id="ref36"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Khan</surname> <given-names>M. S. S.</given-names></name> <name><surname>Rafa</surname> <given-names>S. R.</given-names></name> <name><surname>Das</surname> <given-names>A. K.</given-names></name></person-group> (<year>2021</year>). <article-title>Sentiment analysis on Bengali Facebook comments to predict fan&#x2019;s emotions towards a celebrity</article-title>. <source>J. Eng. Adv.</source> <volume>2</volume>, <fpage>118</fpage>&#x2013;<lpage>124</lpage>.</mixed-citation></ref>
<ref id="ref37"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Kowsher</surname> <given-names>M.</given-names></name> <name><surname>Sami</surname> <given-names>A. A.</given-names></name> <name><surname>Prottasha</surname> <given-names>N. J.</given-names></name> <name><surname>Arefin</surname> <given-names>M. S.</given-names></name> <name><surname>Dhar</surname> <given-names>P. K.</given-names></name> <name><surname>Koshiba</surname> <given-names>T.</given-names></name></person-group> (<year>2022</year>). <article-title>Bangla-BERT: transformer-based efficient model for transfer learning and language understanding</article-title>. <source>IEEE Access</source> <volume>10</volume>, <fpage>91855</fpage>&#x2013;<lpage>91870</lpage>. doi: <pub-id pub-id-type="doi">10.1109/ACCESS.2022.3197662</pub-id>, PMID: <pub-id pub-id-type="pmid">40793584</pub-id></mixed-citation></ref>
<ref id="ref38"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Luhn</surname> <given-names>H. P.</given-names></name></person-group> (<year>1958</year>). <article-title>The automatic creation of literature abstracts</article-title>. <source>IBM J. Res. Dev.</source> <volume>2</volume>, <fpage>159</fpage>&#x2013;<lpage>165</lpage>. doi: <pub-id pub-id-type="doi">10.1147/rd.22.0159</pub-id></mixed-citation></ref>
<ref id="ref39"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Mahmud</surname> <given-names>M. A. I.</given-names></name> <name><surname>Talukder</surname> <given-names>A. T.</given-names></name> <name><surname>Sultana</surname> <given-names>A.</given-names></name> <name><surname>Bhuiyan</surname> <given-names>K. I. A.</given-names></name> <name><surname>Rahman</surname> <given-names>M. S.</given-names></name> <name><surname>Pranto</surname> <given-names>T. H.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Toward news authenticity: synthesizing natural language processing and human expert opinion to evaluate news</article-title>. <source>IEEE Access</source> <volume>11</volume>, <fpage>11405</fpage>&#x2013;<lpage>11421</lpage>. doi: <pub-id pub-id-type="doi">10.1109/ACCESS.2023.3241483</pub-id></mixed-citation></ref>
<ref id="ref40"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Maisha</surname> <given-names>S. J.</given-names></name> <name><surname>Masum</surname> <given-names>A. K. M.</given-names></name> <name><surname>Nafisa</surname> <given-names>N.</given-names></name> <name><surname>Muhammad Masum</surname> <given-names>A. K.</given-names></name></person-group> (<year>2021</year>). <article-title>Supervised machine learning algorithms for sentiment analysis of Bangla newspaper</article-title>. <source>Int. J, Innov. Comp.</source> <volume>11</volume>, <fpage>15</fpage>&#x2013;<lpage>23</lpage>. doi: <pub-id pub-id-type="doi">10.11113/ijic.v11n2.321</pub-id></mixed-citation></ref>
<ref id="ref41"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Mridha</surname> <given-names>M. F.</given-names></name> <name><surname>Wadud</surname> <given-names>M. A. H.</given-names></name> <name><surname>Hamid</surname> <given-names>M. A.</given-names></name> <name><surname>Monowar</surname> <given-names>M. M.</given-names></name> <name><surname>Abdullah-AlWadud</surname> <given-names>M.</given-names></name> <name><surname>Alamri</surname> <given-names>A.</given-names></name></person-group> (<year>2021</year>). <article-title>L-boost: identifying offensive texts from social media post in Bengali</article-title>. <source>IEEE Access</source> <volume>9</volume>, <fpage>164681</fpage>&#x2013;<lpage>164699</lpage>. doi: <pub-id pub-id-type="doi">10.1109/ACCESS.2021.3134154</pub-id></mixed-citation></ref>
<ref id="ref42"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Paice</surname> <given-names>C. D.</given-names></name></person-group> (<year>1990</year>). <article-title>Another stemmer</article-title>. <source>ACM SIGIR Forum</source> <volume>24</volume>, <fpage>56</fpage>&#x2013;<lpage>61</lpage>. doi: <pub-id pub-id-type="doi">10.1145/101306.101310</pub-id></mixed-citation></ref>
<ref id="ref43"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Prottasha</surname> <given-names>N. J.</given-names></name> <name><surname>Sami</surname> <given-names>A. A.</given-names></name> <name><surname>Kowsher</surname> <given-names>M.</given-names></name> <name><surname>Murad</surname> <given-names>S. A.</given-names></name> <name><surname>Bairagi</surname> <given-names>A. K.</given-names></name> <name><surname>Masud</surname> <given-names>M.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Transfer learning for sentiment analysis using BERT based supervised fine-tuning</article-title>. <source>Sensors</source> <volume>22</volume>:<fpage>4157</fpage>. doi: <pub-id pub-id-type="doi">10.3390/s22114157</pub-id>, PMID: <pub-id pub-id-type="pmid">35684778</pub-id></mixed-citation></ref>
<ref id="ref44"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Qiqieh</surname> <given-names>I.</given-names></name> <name><surname>Alzubi</surname> <given-names>O.</given-names></name> <name><surname>Alzubi</surname> <given-names>J.</given-names></name> <name><surname>Sreedhar</surname> <given-names>K. C.</given-names></name> <name><surname>Al-Zoubi</surname> <given-names>A. M.</given-names></name></person-group> (<year>2025</year>). <article-title>An intelligent cyber threat detection: a swarm-optimized machine learning approach</article-title>. <source>Alex. Eng. J.</source> <volume>115</volume>, <fpage>553</fpage>&#x2013;<lpage>563</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.aej.2024.12.039</pub-id></mixed-citation></ref>
<ref id="ref45"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Ramdhani</surname> <given-names>M. A.</given-names></name> <name><surname>Maylawati</surname> <given-names>D. S.</given-names></name> <name><surname>Mantoro</surname> <given-names>T.</given-names></name></person-group> (<year>2020</year>). <article-title>Indonesian news classification using convolutional neural network</article-title>. <source>Indones. J. Electr. Eng. Comput. Sci.</source> <volume>19</volume>, <fpage>1000</fpage>&#x2013;<lpage>1009</lpage>. doi: <pub-id pub-id-type="doi">10.11591/ijeecs.v19.i2.pp1000-1009</pub-id></mixed-citation></ref>
<ref id="ref46"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Rennie</surname> <given-names>J. D.</given-names></name></person-group> (<year>2001</year>). <article-title>Improving multi-class text classification with naive bayes</article-title>.</mixed-citation></ref>
<ref id="ref47"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Ribeiro</surname> <given-names>M. T.</given-names></name> <name><surname>Singh</surname> <given-names>S.</given-names></name> <name><surname>Guestrin</surname> <given-names>C.</given-names></name></person-group> (<year>2016</year>). &#x201C;Why should I trust you?&#x2019; Explaining the predictions of any classifier.&#x201D; In <italic>Proceedings of the 22nd ACM SIGKDD International Conference on Knowledge Discovery and Data Mining</italic>. pp. 1135&#x2013;1144.</mixed-citation></ref>
<ref id="ref48"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Roy</surname> <given-names>A.</given-names></name> <name><surname>Sarkar</surname> <given-names>K.</given-names></name> <name><surname>Mandal</surname> <given-names>C. K.</given-names></name></person-group> (<year>2023</year>). <source>Bengali text classification: A new multiclass dataset and performance evaluation of machine learning and deep learning models</source>. <publisher-loc>Amsterdam, Netherlands</publisher-loc>: <publisher-name>Elsevier</publisher-name>.</mixed-citation></ref>
<ref id="ref49"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Saigal</surname> <given-names>P.</given-names></name> <name><surname>Khanna</surname> <given-names>V.</given-names></name></person-group> (<year>2020</year>). <article-title>Multi-category news classification using support vector machine-based classifiers</article-title>. <source>SN Appl. Sci.</source> <volume>2</volume>:<fpage>458</fpage>. doi: <pub-id pub-id-type="doi">10.1007/s42452-020-2266-6</pub-id></mixed-citation></ref>
<ref id="ref50"><mixed-citation publication-type="book"><person-group person-group-type="author"><name><surname>Salton</surname> <given-names>G.</given-names></name></person-group> (<year>1983</year>). <source>Introduction to modern information retrieval</source>. <publisher-loc>New York, NY, USA</publisher-loc>: <publisher-name>McGrawHill Book Co.</publisher-name></mixed-citation></ref>
<ref id="ref51"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Sanh</surname> <given-names>V.</given-names></name></person-group> (<year>2019</year>). <article-title>DistilBERT, a distilled version of BERT: smaller, faster, cheaper and lighter</article-title>. <source>arXiv</source>:<fpage>1910.01108</fpage>.</mixed-citation></ref>
<ref id="ref52"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Sen</surname> <given-names>O.</given-names></name> <name><surname>Fuad</surname> <given-names>M.</given-names></name> <name><surname>Islam</surname> <given-names>M. N.</given-names></name> <name><surname>Rabbi</surname> <given-names>J.</given-names></name> <name><surname>Masud</surname> <given-names>M.</given-names></name> <name><surname>Hasan</surname> <given-names>M. K.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Bangla natural language processing: a comprehensive analysis of classical, machine learning, and deep learning-based methods</article-title>. <source>IEEE Access</source> <volume>10</volume>, <fpage>38999</fpage>&#x2013;<lpage>39044</lpage>. doi: <pub-id pub-id-type="doi">10.1109/ACCESS.2022.3165563</pub-id></mixed-citation></ref>
<ref id="ref53"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Sourav</surname> <given-names>M. S. U.</given-names></name> <name><surname>Wang</surname> <given-names>H.</given-names></name> <name><surname>Mahmud</surname> <given-names>M. S.</given-names></name> <name><surname>Zheng</surname> <given-names>H.</given-names></name></person-group> (<year>2022</year>). <article-title>Transformer-based text classification on unified Bangla multi-class emotion corpus</article-title>. <source>arXiv</source>:<fpage>2210.06405</fpage>.</mixed-citation></ref>
<ref id="ref54"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Tareq</surname> <given-names>M.</given-names></name> <name><surname>Islam</surname> <given-names>M. F.</given-names></name> <name><surname>Deb</surname> <given-names>S.</given-names></name> <name><surname>Rahman</surname> <given-names>S.</given-names></name> <name><surname>Al Mahmud</surname> <given-names>A.</given-names></name></person-group> (<year>2023</year>). <article-title>Data augmentation for Bangla-English code-mixed sentiment analysis: enhancing cross-linguistic contextual understanding</article-title>. <source>IEEE Access</source> <volume>11</volume>, <fpage>51657</fpage>&#x2013;<lpage>51671</lpage>. doi: <pub-id pub-id-type="doi">10.1109/ACCESS.2023.3277787</pub-id></mixed-citation></ref>
<ref id="ref55"><mixed-citation publication-type="other"><person-group person-group-type="author"><name><surname>Timeline</surname> <given-names>B.</given-names></name></person-group> (<year>n.d.</year>). All Bangla newspapers. Available online at: <ext-link xlink:href="https://www.allbanglanewspaper.xyz/" ext-link-type="uri">https://www.allbanglanewspaper.xyz/</ext-link> (Accessed July 2, 2023).</mixed-citation></ref>
<ref id="ref56"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Venkatsubramaniam</surname> <given-names>B.</given-names></name> <name><surname>Baruah</surname> <given-names>P. K.</given-names></name></person-group> (<year>2022</year>). <article-title>Comparative study of XAI using formal concept lattice and LIME</article-title>. <source>ICTACT J. Soft Comp.</source> <volume>13</volume>, <fpage>2782</fpage>&#x2013;<lpage>2791</lpage>. doi: <pub-id pub-id-type="doi">10.21917/ijsc.2022.0396</pub-id></mixed-citation></ref>
<ref id="ref57"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Wadud</surname> <given-names>M. A. H.</given-names></name> <name><surname>Mridha</surname> <given-names>M.</given-names></name> <name><surname>Rahman</surname> <given-names>M. M.</given-names></name></person-group> (<year>2022</year>). <article-title>Word embedding methods for word representation in deep learning for natural language processing</article-title>. <source>Iraqi J. Sci.</source>, <volume>63</volume>, <fpage>1349</fpage>&#x2013;<lpage>1361</lpage>. doi: <pub-id pub-id-type="doi">10.24996/ijs.2022.63.3.37</pub-id></mixed-citation></ref>
<ref id="ref58"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Yeasmin</surname> <given-names>S.</given-names></name> <name><surname>Kuri</surname> <given-names>R.</given-names></name> <name><surname>Rana</surname> <given-names>A.</given-names></name> <name><surname>Uddin</surname> <given-names>A.</given-names></name> <name><surname>Pathan</surname> <given-names>A.</given-names></name> <name><surname>Riaz</surname> <given-names>H.</given-names></name></person-group> (<year>2021</year>). <article-title>Multi-category Bangla news classification using machine learning classifiers and multi-layer dense neural network</article-title>. <source>Int. J. Adv. Comput. Sci. Appl.</source> <volume>12</volume>:<fpage>5</fpage>. doi: <pub-id pub-id-type="doi">10.14569/IJACSA.2021.0120588</pub-id></mixed-citation></ref>
<ref id="ref59"><mixed-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>M.</given-names></name></person-group> (<year>2021</year>). <article-title>Applications of deep learning in news text classification</article-title>. <source>Sci. Program.</source> <volume>2021</volume>:<fpage>6095354</fpage>.</mixed-citation></ref>
</ref-list><fn-group><fn id="fn0001" fn-type="custom" custom-type="edited-by"><p>Edited by: <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1859724/overview">P. K. Gupta</ext-link>, Jaypee University of Information Technology, India</p></fn>
<fn id="fn0002" fn-type="custom" custom-type="reviewed-by"><p>Reviewed by: <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1858648/overview">Omar A. Alzubi</ext-link>, Al-Balqa Applied University, Jordan; <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2348836/overview">Sumali Conlon</ext-link>, University of Mississippi, United States; <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2413280/overview">Osman Ali Sadek Ibrahim</ext-link>, Minia University, Egypt</p></fn></fn-group></back>
</article>
