<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Archiving and Interchange DTD v2.3 20070202//EN" "archivearticle.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="systematic-review" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Oncol.</journal-id>
<journal-title>Frontiers in Oncology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Oncol.</abbrev-journal-title>
<issn pub-type="epub">2234-943X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fonc.2022.856231</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Oncology</subject>
<subj-group>
<subject>Systematic Review</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Machine Learning Models for Classifying High- and Low-Grade Gliomas: A Systematic Review and Quality of Reporting Analysis</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Bahar</surname>
<given-names>Ryan C.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="author-notes" rid="fn003">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1622814"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Merkaj</surname>
<given-names>Sara</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="author-notes" rid="fn003">
<sup>&#x2020;</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1635071"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Cassinelli Petersen</surname>
<given-names>Gabriel I.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1503298"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Tillmanns</surname>
<given-names>Niklas</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1503253"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Subramanian</surname>
<given-names>Harry</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1452486"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Brim</surname>
<given-names>Waverly Rose</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Zeevi</surname>
<given-names>Tal</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Staib</surname>
<given-names>Lawrence</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Kazarian</surname>
<given-names>Eve</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1716112"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Lin</surname>
<given-names>MingDe</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Bousabarah</surname>
<given-names>Khaled</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Huttner</surname>
<given-names>Anita J.</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1716536"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Pala</surname>
<given-names>Andrej</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/411149"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Payabvash</surname>
<given-names>Seyedmehdi</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/611713"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Ivanidze</surname>
<given-names>Jana</given-names>
</name>
<xref ref-type="aff" rid="aff6">
<sup>6</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/679237"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Cui</surname>
<given-names>Jin</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Malhotra</surname>
<given-names>Ajay</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Aboian</surname>
<given-names>Mariam S.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>*</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1647057"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Department of Radiology and Biomedical Imaging, Yale School of Medicine</institution>, <addr-line>New Haven, CT</addr-line>, <country>United States</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Department of Neurosurgery, University of Ulm</institution>, <addr-line>Ulm</addr-line>, <country>Germany</country>
</aff>
<aff id="aff3">
<sup>3</sup>
<institution>Visage Imaging, Inc.</institution>, <addr-line>San Diego, CA</addr-line>, <country>United States</country>
</aff>
<aff id="aff4">
<sup>4</sup>
<institution>Visage Imaging</institution>, <addr-line>GmbH., Berlin</addr-line>, <country>Germany</country>
</aff>
<aff id="aff5">
<sup>5</sup>
<institution>Department of Pathology, Yale-New Haven Hospital, Yale School of Medicine</institution>, <addr-line>New Haven, CT</addr-line>, <country>United States</country>
</aff>
<aff id="aff6">
<sup>6</sup>
<institution>Department of Radiology, Weill Cornell Medicine</institution>, <addr-line>New York, NY</addr-line>, <country>United States</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>Edited by: Dario de Biase, University of Bologna, Italy</p>
</fn>
<fn fn-type="edited-by">
<p>Reviewed by: Vincenzo Di Nunno, AUSL Bologna, Italy; Chirag Kamal Ahuja, Post Graduate Institute of Medical Education and Research (PGIMER), India</p>
</fn>
<fn fn-type="corresp" id="fn001">
<p>*Correspondence: Mariam S. Aboian, <email xlink:href="mailto:mariam.aboian@yale.edu">mariam.aboian@yale.edu</email>
</p>
</fn>
<fn fn-type="other" id="fn002">
<p>This article was submitted to Neuro-Oncology and Neurosurgical Oncology, a section of the journal Frontiers in Oncology</p>
</fn>
<fn fn-type="equal" id="fn003">
<p>&#x2020;These authors have contributed equally to this work and share first authorship</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>22</day>
<month>04</month>
<year>2022</year>
</pub-date>
<pub-date pub-type="collection">
<year>2022</year>
</pub-date>
<volume>12</volume>
<elocation-id>856231</elocation-id>
<history>
<date date-type="received">
<day>16</day>
<month>01</month>
<year>2022</year>
</date>
<date date-type="accepted">
<day>25</day>
<month>03</month>
<year>2022</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2022 Bahar, Merkaj, Cassinelli Petersen, Tillmanns, Subramanian, Brim, Zeevi, Staib, Kazarian, Lin, Bousabarah, Huttner, Pala, Payabvash, Ivanidze, Cui, Malhotra and Aboian</copyright-statement>
<copyright-year>2022</copyright-year>
<copyright-holder>Bahar, Merkaj, Cassinelli Petersen, Tillmanns, Subramanian, Brim, Zeevi, Staib, Kazarian, Lin, Bousabarah, Huttner, Pala, Payabvash, Ivanidze, Cui, Malhotra and Aboian</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<sec>
<title>Objectives</title>
<p>To systematically review, assess the reporting quality of, and discuss improvement opportunities for studies describing machine learning (ML) models for glioma grade prediction.</p>
</sec>
<sec>
<title>Methods</title>
<p>This study followed the Preferred Reporting Items for Systematic Reviews and Meta-Analyses of Diagnostic Test Accuracy (PRISMA-DTA) statement. A systematic search was performed in September 2020, and repeated in January 2021, on four databases: Embase, Medline, CENTRAL, and Web of Science Core Collection. Publications were screened in Covidence, and reporting quality was measured against the Transparent Reporting of a multivariable prediction model for Individual Prognosis Or Diagnosis (TRIPOD) Statement. Descriptive statistics were calculated using GraphPad Prism 9.</p>
</sec>
<sec>
<title>Results</title>
<p>The search identified 11,727 candidate articles with 1,135 articles undergoing full text review and 85 included in analysis. 67 (79%) articles were published between 2018-2021. The mean prediction accuracy of the best performing model in each study was 0.89 &#xb1; 0.09. The most common algorithm for conventional machine learning studies was Support Vector Machine (mean accuracy: 0.90 &#xb1; 0.07) and for deep learning studies was Convolutional Neural Network (mean accuracy: 0.91 &#xb1; 0.10). Only one study used both a large training dataset (n&gt;200) and external validation (accuracy: 0.72) for their model. The mean adherence rate to TRIPOD was 44.5% &#xb1; 11.1%, with poor reporting adherence for model performance (0%), abstracts (0%), and titles (0%).</p>
</sec>
<sec>
<title>Conclusions</title>
<p>The application of ML to glioma grade prediction has grown substantially, with ML model studies reporting high predictive accuracies but lacking essential metrics and characteristics for assessing model performance. Several domains, including generalizability and reproducibility, warrant further attention to enable translation into clinical practice.</p>
</sec>
<sec>
<title>Systematic Review Registration</title>
<p>PROSPERO, identifier CRD42020209938.</p>
</sec>
</abstract>
<kwd-group>
<kwd>machine learning</kwd>
<kwd>deep learning</kwd>
<kwd>artificial intelligence</kwd>
<kwd>glioma</kwd>
<kwd>systematic review</kwd>
</kwd-group>
<contract-num rid="cn001">BMEP</contract-num>
<contract-num rid="cn002">T35DK104689</contract-num>
<contract-num rid="cn003">Fellow Award 2018</contract-num>
<contract-num rid="cn004">KL2 TR001862</contract-num>
<contract-num rid="cn005">NCI R01 CA206180</contract-num>
<contract-num rid="cn006">NINDS K23NS118056</contract-num>
<contract-num rid="cn007">1861150721</contract-num>
<contract-num rid="cn008">2020097</contract-num>
<contract-sponsor id="cn001">Deutscher Akademischer Austauschdienst<named-content content-type="fundref-id">10.13039/501100001655</named-content>
</contract-sponsor>
<contract-sponsor id="cn002">National Institute of Diabetes and Digestive and Kidney Diseases<named-content content-type="fundref-id">10.13039/100000062</named-content>
</contract-sponsor>
<contract-sponsor id="cn003">American Society of Neuroradiology<named-content content-type="fundref-id">10.13039/100011975</named-content>
</contract-sponsor>
<contract-sponsor id="cn004">National Center for Advancing Translational Sciences<named-content content-type="fundref-id">10.13039/100006108</named-content>
</contract-sponsor>
<contract-sponsor id="cn005">National Institutes of Health<named-content content-type="fundref-id">10.13039/100000002</named-content>
</contract-sponsor>
<contract-sponsor id="cn006">National Institutes of Health<named-content content-type="fundref-id">10.13039/100000002</named-content>
</contract-sponsor>
<contract-sponsor id="cn007">American Society of Neuroradiology<named-content content-type="fundref-id">10.13039/100011975</named-content>
</contract-sponsor>
<contract-sponsor id="cn008">Doris Duke Charitable Foundation<named-content content-type="fundref-id">10.13039/100000862</named-content>
</contract-sponsor>
<contract-sponsor id="cn009">Nvidia<named-content content-type="fundref-id">10.13039/100007065</named-content>
</contract-sponsor>
<counts>
<fig-count count="6"/>
<table-count count="2"/>
<equation-count count="0"/>
<ref-count count="55"/>
<page-count count="11"/>
<word-count count="4768"/>
</counts>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>Gliomas are the most common primary brain malignancy (<xref ref-type="bibr" rid="B1">1</xref>). They are classified according to histopathologic and molecular World Health Organization (WHO) criteria: grades 1/2 (low-grade gliomas (LGG) and grades 3/4 [high-grade gliomas (HGG)] (<xref ref-type="bibr" rid="B2">2</xref>). Glioblastomas, WHO grade 4 tumors, are the most aggressive with a 15-month median overall survival (<xref ref-type="bibr" rid="B3">3</xref>).</p>
<p>Because prognosis (<xref ref-type="bibr" rid="B3">3</xref>, <xref ref-type="bibr" rid="B4">4</xref>) and treatment (<xref ref-type="bibr" rid="B5">5</xref>) vary with glioma grade, accurate classification is essential for guiding clinical decision-making and mitigating risks posed by unnecessary or delayed surgery due to misdiagnosis (<xref ref-type="bibr" rid="B6">6</xref>). The gold standard for diagnosis, histopathology, requires surgical resection or stereotactic biopsy for analysis. These invasive procedures, however, carry significant risks and complications (<xref ref-type="bibr" rid="B7">7</xref>). Gliomas also exhibit intratumoral heterogeneity with associated sampling error (<xref ref-type="bibr" rid="B8">8</xref>). Therefore, a need exists for timely pre-operative whole-glioma grading. As a non-invasive tool for analyzing entire lesions, imaging overcomes the limitations of diagnostic surgical procedures. Although conventional MRI has had modest success in glioma grading (sensitivity 55-83%) (<xref ref-type="bibr" rid="B9">9</xref>), the diagnostic potential of imaging has expanded with the use of advanced imaging, radiomics, and artificial intelligence.</p>
<p>Radiomics quantitatively characterizes medical images using image-derived features that serve as biomarkers for tumor phenotypes (<xref ref-type="bibr" rid="B10">10</xref>). Artificial intelligence technologies, such as machine learning (ML), have augmented radiomics. By leveraging robust high-dimensional data, ML enhances predictive performance (<xref ref-type="bibr" rid="B11">11</xref>). Deep learning (DL) is a subtype of ML that has sparked recent interest given its superior performance in image analysis and suitability for high volumes of data (<xref ref-type="bibr" rid="B12">12</xref>). For imaging applications, DL generates useful outputs from input images using multilayer neural networks. Convolutional Neural Networks are the primary DL architecture for image classification (<xref ref-type="bibr" rid="B13">13</xref>).</p>
<p>In clinical practice, ML models may increase the value of diagnostic imaging and enhance patient management, for example, by motivating earlier grade-appropriate interventions (<xref ref-type="bibr" rid="B14">14</xref>, <xref ref-type="bibr" rid="B15">15</xref>). Despite these opportunities, ML has not been implemented clinically because of numerous technical (data requirements, need for training, low standardization and interpretability) and non-technical (ethical, financial, legal, educational) barriers (<xref ref-type="bibr" rid="B16">16</xref>).</p>
<p>High-quality scientific reporting is necessary for readers to critically interpret or replicate studies and encourage translation into practice. Prior work indicates that reporting quality in prediction studies is poor (<xref ref-type="bibr" rid="B17">17</xref>). To address this, the Transparent Reporting of a multivariable model for Individual Prognosis or Diagnosis (TRIPOD) Statement was published in 2015 (<xref ref-type="bibr" rid="B18">18</xref>). Most of TRIPOD is applicable to ML-based prediction model studies; however, ML-specific guidelines are lacking. The need for such guidelines has initiated development of a TRIPOD extension for ML-based prediction models (TRIPOD-AI) (<xref ref-type="bibr" rid="B19">19</xref>, <xref ref-type="bibr" rid="B20">20</xref>).</p>
<p>While ML demonstrates promise for accurate glioma grading, few works have characterized the state of ML in glioma grade prediction (<xref ref-type="bibr" rid="B21">21</xref>&#x2013;<xref ref-type="bibr" rid="B23">23</xref>). A systematic review of the literature can identify potential ML methods for clinical use and generate insights for implementation. This study aims to (1) systematically review and synthesize the body of literature using ML for classification of glioma grade, (2) evaluate study reporting quality using TRIPOD, and (3) discuss opportunities for bridging the ML bench-to-clinic implementation gap.</p>
</sec>
<sec id="s2" sec-type="materials|methods">
<title>2 Materials and Methods</title>
<p>This study followed the guidelines in the Preferred Reporting Items for Systematic Reviews and Meta-Analyses of Diagnostic Test Accuracy (PRISMA-DTA) statement (<xref ref-type="bibr" rid="B24">24</xref>) and was registered with the International Prospective Register of Systematic Reviews (PROSPERO, CRD42020209938). An institutional librarian searched the literature published through September 18, 2020, using four databases: Cochrane Central Register of Controlled Trials, EMBASE, Medline, and Web of Science Core Collection. A multi-database approach was pursued because prior work has demonstrated that a single database search may omit pertinent studies (<xref ref-type="bibr" rid="B25">25</xref>). Keywords and controlled vocabulary included the following terms and combinations thereof: &#x201c;artificial intelligence,&#x201d; &#x201c;machine learning,&#x201d; &#x201c;deep learning,&#x201d; &#x201c;radiomics,&#x201d; &#x201c;magnetic resonance imaging,&#x201d; &#x201c;glioma,&#x201d; and related terms. The search was repeated on January 29, 2021, to gather additional articles published through this date. A full search strategy is provided in <xref ref-type="supplementary-material" rid="ST1">
<bold>Appendix A1 (Supplementary)</bold>
</xref>. A second institutional librarian reviewed the search prior to execution.</p>
<p>Only peer-reviewed studies were imported into Covidence (Veritas Health Innovation Ltd) for screening. Covidence is an online tool designed to streamline the systematic review process. Duplicate studies were identified and removed. Study abstracts were then screened for relevance to neuro-oncology by two of three independent reviewers: an experienced board-certified neuroradiologist, radiology resident, and graduate student in artificial intelligence. The board-certified neuroradiologist resolved discrepancies in screening recommendations. Relevant articles were subsequently assessed for eligibility. To ensure completeness, appropriateness, and understandability of eligible studies, the following exclusion criteria were established: (1) abstract-only; (2) not primary literature; (3) non-English; (4) unrelated to artificial intelligence; (5) unrelated to gliomas; (6) unrelated to imaging; (7) non-human research subjects; and (8) duplicates. Studies were not excluded based on publication year in order to have comprehensive analysis of historical and contemporary literature. Eligible studies underwent full text review to identify those using ML to classify gliomas by grade. Studies exclusively developing predictive models with distinct focuses (e.g., predicting glioma IDH status, glioma segmentation) were not included in analysis. A PRISMA flow diagram describing our study selection process is depicted in <xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref>.</p>
<fig id="f1" position="float">
<label>Figure&#xa0;1</label>
<caption>
<p>PRISMA flow diagram of study search strategy. ML, machine learning; PRISMA, Preferred Reporting Guidelines for Systematic Reviews and Meta-Analyses.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fonc-12-856231-g001.tif"/>
</fig>
<p>Whole data was independently extracted by two trained medical student researchers using a standardized Microsoft Excel (Microsoft Corporation) form. Conflicts were resolved through team discussion and consensus. When necessary, articles were carefully re-reviewed to obtain missing information after data extraction. The following data points were extracted: article characteristics (title, lead author, country of lead author, publication year), data characteristics (data source, country (or countries) of data acquisition, dataset size, types and number of tumors for training/testing/validation, model validation technique), grading characteristics (study definition of HGG and LGG, gold standard for glioma grading), model characteristics (best performing ML classifier, classification task, supervised/semi-supervised/unsupervised learning, types of features in classifier, imaging sequences used by classifier, measures of classifier performance) and reporting characteristics (TRIPOD items, explained below).</p>
<p>Reporting quality was assessed against the TRIPOD statement in agreement with the TRIPOD adherence assessment form (<xref ref-type="bibr" rid="B26">26</xref>) and author explanations (<xref ref-type="bibr" rid="B18">18</xref>, <xref ref-type="bibr" rid="B27">27</xref>). TRIPOD contains 20 main items (e.g., main item 5) that apply to studies developing prediction models, 10 of which contain subitems (e.g., 5a, 5b, 5c). Among the 30 total items that can be evaluated and scored, three (item 5c, 11, 14b) were excluded because they were not applicable to our studies. The remaining 27 items were scored for every study. Each item includes one or more elements, all of which must score a &#x201c;yes&#x201d; for the item to score &#x201c;1.&#x201d; To calculate a study&#x2019;s adherence rate to TRIPOD, the number of items scoring &#x201c;1&#x201d; was divided by the total number of scored items for the study. Adherence rate for a given TRIPOD item across all studies was calculated by dividing the number of studies scoring &#x201c;1&#x201d; for that item by the total number of studies scored.</p>
<p>TRIPOD adherence rates and descriptive statistics (e.g., frequencies, mean &#xb1; standard deviation) were calculated and displayed with GraphPad Prism 9 (GraphPad Software). GraphPad Prism 9 is a scientific graphing and statistical software supporting data analysis. Descriptive statistics were obtained to summarize study characteristics, dataset and model characteristics, features, imaging modalities, and model prediction performance, among other domains. Only the best performing classifiers&#x2019; performance metrics (accuracy, AUC, sensitivity, specificity, positive predictive value, negative predictive value, and F1 score) defined in each study are presented. Best performing classifiers were determined based on accuracy results. In the few instances when accuracy was not reported, AUC determined the best performing classifier. All studies meeting our inclusion criteria ((1) identified by search strategy; (2) relevant to neuro-oncology; (3) not excluded during eligibility assessment; and (4) clearly evaluate ML-based classification of glioma grade) contributed to the sample size of our study. References for included studies are listed in <xref ref-type="supplementary-material" rid="ST1">
<bold>Appendix A4 (Supplementary)</bold>
</xref>.</p>
</sec>
<sec id="s3">
<title>3 Results</title>
<sec id="s3_1">
<title>3.1 Study Characteristics</title>
<p>The search identified 11,727 candidate articles, with 11,637 studies screened for relevance to neuro-oncology. Agreement between screeners was substantial [(Cohen&#x2019;s kappa: 0.77 &#xb1; 0.04, see <xref ref-type="supplementary-material" rid="ST1">
<bold>Table A1 (Supplementary)</bold>
</xref>]. 1,135 articles underwent full text review, and 85 articles were included in analysis (<xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref>).</p>
<p>67 articles (79%) were published between 2018 and 2021, with 26 articles (31%) published in 2019 alone (<xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>). Based on lead author affiliations, most articles were from China, the US, or India (n=45, 51%) (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref>).</p>
<fig id="f2" position="float">
<label>Figure&#xa0;2</label>
<caption>
<p>Number of studies published per year from 1995-2020.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fonc-12-856231-g002.tif"/>
</fig>
<fig id="f3" position="float">
<label>Figure&#xa0;3</label>
<caption>
<p>
<bold>(A)</bold> Number of studies by first author&#x2019;s country of affiliation and respective continent. <bold>(B)</bold> Number of studies by country (or countries) of data acquisition.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fonc-12-856231-g003.tif"/>
</fig>
<p>36 articles (42%) defined HGG as grade 3 and 4 and LGG as grade 1 and 2. 17 articles (20%) defined HGG as grade 4 and LGG as grades 2 and 3. 32 articles (38%) didn&#x2019;t define grades for HGG and LGG (<xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4</bold>
</xref>).</p>
<fig id="f4" position="float">
<label>Figure&#xa0;4</label>
<caption>
<p>Classification systems used across studies for defining HGG vs. LGG by grade 1-4. HGG, high-grade gliomas; LGG, low-grade gliomas.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fonc-12-856231-g004.tif"/>
</fig>
</sec>
<sec id="s3_2">
<title>3.2 Study Findings</title>
<sec id="s3_2_1">
<title>3.2.1 Dataset and Model Characteristics</title>
<p>Among the 84 articles with identifiable patient data sources, data was most commonly acquired multi-nationally (n=38, 45%), entirely in China (n=15, 18%), or entirely in the US (n=10, 12%) (<xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref>). BraTS (<xref ref-type="bibr" rid="B28">28</xref>) and TCIA (<xref ref-type="bibr" rid="B29">29</xref>) datasets, which are publicly available multi-institutional datasets containing multi-parametric MRI scans, were used in 45% (n=38) of studies. Conventional ML was the primary ML model for glioma grade prediction in 59 (69%) studies and DL in 26 (31%) studies. Of all 85 studies included, 80 (94%) reported the number of patients in their datasets (mean: 177 &#xb1; 140). Studies developing conventional ML models reported mean dataset sizes of 168 &#xb1; 150 patients. DL studies reported mean dataset sizes of 199 &#xb1; 109 patients. Among the 67 studies whose best performing models were binary classifiers of HGG and LGG and reported the number of HGG and LGG used in model development, 58 (87%) had imbalanced datasets characterized by an unequal number of HGG and LGG patients. Most studies (n=44, 66%) used datasets containing more HGG than LGG patients (i.e., HGG : LGG ratio &gt;1). 14 studies (21%) had fewer HGG than LGG patients (HGG : LGG ratio &lt;1).</p>
<p>Only 5 (6%) studies reported external validation. Of the 80 other studies, 68 (85%) reported internal validation and 12 (15%) did not clearly report validation methods. 82 (96%) of studies had supervised learning algorithms and 3 (4%) used semi-supervised learning. No studies reported unsupervised learning algorithms. The gold standard for glioma grading was histopathology in all studies.</p>
</sec>
<sec id="s3_2_2">
<title>3.2.2 Features</title>
<p>Texture (second-order) features and first-order features were the most common feature subsets, extracted in 45 (53%) and 42 (49%) studies, respectively. Shape and/or size features (n=28, 33%) and DL extracted features (n=20, 24%) were also common. Hemodynamic (n=5), qualitative (n=6), higher-order (n=4) and spectroscopic features (n=8) were observed in less than 10% of studies. Definitions for feature types are provided in <xref ref-type="supplementary-material" rid="ST1">
<bold>Table A7 (Supplementary)</bold>
</xref>.</p>
</sec>
<sec id="s3_2_3">
<title>3.2.3 Imaging Modalities</title>
<p>T1-weighted contrast-enhanced (T1CE) imaging was the most common sequence used in best performing models (n=54, 64%), followed by T2 (n=46, 54%) and FLAIR (n=40, 47%). T1 pre-contrast was less common (n=35, 41%). Perfusion-weighted imaging (n=15), MR Spectroscopy (n=9) and diffusion-weighted imaging (n=12) were used in 11-18% of models. PET and fMRI were only used in one model each.</p>
</sec>
<sec id="s3_2_4">
<title>3.2.4 Prediction Performance</title>
<p>A summary of model performance measures across studies is shown in <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref>. The mean glioma grade prediction accuracy of the best performing algorithm per study was 0.89 &#xb1; 0.09. This parameter was determined by taking the prediction accuracy of the best performing algorithm in each study for all studies and calculating a mean value and standard deviation. Lower accuracies were reported for models undergoing external validation (mean: 0.82 &#xb1; 0.09, n=5). DL models had a mean prediction accuracy of 0.92 &#xb1; 0.08 and conventional ML models 0.88 &#xb1; 0.09.</p>
<table-wrap id="T1" position="float">
<label>Table&#xa0;1</label>
<caption>
<p>Mean (&#xb1; standard deviation) aggregate performance metrics across studies.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="left">Accuracy (n=82)</th>
<th valign="top" align="center">AUC (n=48)</th>
<th valign="top" align="center">Sensitivity (n=55)</th>
<th valign="top" align="center">Specificity (n=51)</th>
<th valign="top" align="center">Positive Predictive Value (n=12)</th>
<th valign="top" align="center">Negative Predictive Value (n=6)</th>
<th valign="top" align="center">F1 Score (n=7)</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">0.89 &#xb1; 0.09</td>
<td valign="top" align="left">0.92 &#xb1; 0.07</td>
<td valign="top" align="center">0.89 &#xb1; 0.09</td>
<td valign="top" align="center">0.88 &#xb1; 0.11</td>
<td valign="top" align="center">0.90 &#xb1; 0.09</td>
<td valign="top" align="center">0.82 &#xb1; 0.08</td>
<td valign="top" align="center">0.89 &#xb1; 0.11</td>
</tr>
<tr>
<td valign="top" align="left">(0.53-1.00)</td>
<td valign="top" align="left"> (0.73-1.00)</td>
<td valign="top" align="center">(0.63-1.00)</td>
<td valign="top" align="center">(0.55-1.00)</td>
<td valign="top" align="center">(0.68-1.00)</td>
<td valign="top" align="center">(0.73-0.94)</td>
<td valign="top" align="center">(0.67-0.98)</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>n, number of studies reporting metric.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>The most common best performing conventional ML model was Support Vector Machine (mean accuracy: 0.90 &#xb1; 0.07) and DL model was Convolutional Neural Network (mean accuracy: 0.91 &#xb1; 0.10) (<xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5</bold>
</xref>).</p>
<fig id="f5" position="float">
<label>Figure&#xa0;5</label>
<caption>
<p>Prediction accuracy of most common algorithm types, measured in the best performing algorithm of each study. Circle at mean. Error bars indicate standard deviation.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fonc-12-856231-g005.tif"/>
</fig>
<p>We grouped all studies by data source into 4 categories: BraTS, TCIA, single center, and multicenter (excluding BraTS and TCIA) data. Studies which used BraTS as a data source had a mean accuracy of 0.93 &#xb1; 0.04 (n=27) and studies using TCIA had a mean accuracy of 0.91 &#xb1; 0.08 (n=12). Single center datasets were the most common (n=43) with a mean accuracy of 0.88 &#xb1; 0.07, and multicenter hospital datasets the least common (n=6, mean accuracy: 0.80 &#xb1; 0.18).</p>
<p>We additionally identified studies whose models were built on relatively large (n&#x2265;200) datasets and externally validated, two characteristics indicating potential generalizability. Only one study (1%) had both characteristics (accuracy: 0.72) (<xref ref-type="bibr" rid="B30">30</xref>). Further analysis of model performance by dataset source, dataset size, validation technique, and glioma grade classification task can be found in <xref ref-type="supplementary-material" rid="ST1">
<bold>Appendix A2</bold>
</xref> and <xref ref-type="supplementary-material" rid="ST1">
<bold>Tables A2-A5 (Supplementary)</bold>
</xref>. Characteristics of the 10 studies reporting the highest accuracy results for their best performing algorithms are summarized in <xref ref-type="table" rid="T2">
<bold>Table&#xa0;2</bold>
</xref>. Characteristics of all included studies may be seen in <xref ref-type="supplementary-material" rid="ST1">
<bold>Table A6 (Supplementary)</bold>
</xref>.</p>
<table-wrap id="T2" position="float">
<label>Table&#xa0;2</label>
<caption>
<p>Characteristics of the 10 studies reporting the highest accuracy results for their best performing models, including: glioma grade classification task, dataset source and size, ratio of high- to low-grade gliomas, validation technique, imaging sequences used in prediction, feature types used in prediction, best performing algorithm (based on accuracy results), and performance metrics.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="top" align="left">Paper</th>
<th valign="top" align="center">Glioma Grade Classification Task</th>
<th valign="top" align="center">Dataset</th>
<th valign="top" align="center">HGG : LGG Ratio</th>
<th valign="top" align="center">Validation Technique</th>
<th valign="top" align="center">Imaging Sequences</th>
<th valign="top" align="center">Features</th>
<th valign="top" align="center">Best Algorithm</th>
<th valign="top" align="center">Performance Metrics</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Hedyehzadeh et&#xa0;al. (2020) (<xref ref-type="bibr" rid="B31">31</xref>)</td>
<td valign="top" align="left">2/3 vs. 4</td>
<td valign="top" align="left">TCIA (n=461 patients)</td>
<td valign="top" align="left">1.3:1 (262 HGG, 199 LGG in total set)</td>
<td valign="top" align="left">Internal (4-fold cross-validation)</td>
<td valign="top" align="left">T1, T1CE, T2, FLAIR</td>
<td valign="top" align="left">Texture</td>
<td valign="top" align="left">Support Vector Machine</td>
<td valign="top" align="left">Accuracy = 1.00 Sensitivity = 1.00 Specificity = 1.00</td>
</tr>
<tr>
<td valign="top" align="left">BashirGonbadi and Khotanlou (2019) (<xref ref-type="bibr" rid="B32">32</xref>)</td>
<td valign="top" align="left">1/2 vs. 3/4</td>
<td valign="top" align="left">BraTS (n=285 patients)</td>
<td valign="top" align="left">2.8:1 (210 HGG, 75 LGG in total set)</td>
<td valign="top" align="left">Internal (Holdout, 15% of dataset)</td>
<td valign="top" align="left">T1, T1CE, T2, FLAIR</td>
<td valign="top" align="left">Deep learning extracted</td>
<td valign="top" align="left">Convolutional Neural Network</td>
<td valign="top" align="left">Accuracy = 0.9918</td>
</tr>
<tr>
<td valign="top" align="left">Polly et&#xa0;al. (2018) (<xref ref-type="bibr" rid="B33">33</xref>)</td>
<td valign="top" align="left">HGG vs. LGG (unclear)</td>
<td valign="top" align="left">BraTS (n=160 images)</td>
<td valign="top" align="left">1:1 (50 HGG, 50 LGG in testing set)</td>
<td valign="top" align="left">Unspecified</td>
<td valign="top" align="left">T2</td>
<td valign="top" align="left">First-order, Shape, Texture</td>
<td valign="top" align="left">Support Vector Machine</td>
<td valign="top" align="left">Accuracy = 0.99 Sensitivity = 1.00 Specificity = 0.9803</td>
</tr>
<tr>
<td valign="top" align="left">De Looze et&#xa0;al. (2018) (<xref ref-type="bibr" rid="B34">34</xref>)</td>
<td valign="top" align="left">HGG vs. LGG (unclear)</td>
<td valign="top" align="left">Single center hospital (n=381 patients)</td>
<td valign="top" align="left">Unclear</td>
<td valign="top" align="left">Internal (5-fold cross-validation)</td>
<td valign="top" align="left">T1, T1CE, T2, FLAIR, Diffusion</td>
<td valign="top" align="left">Qualitative</td>
<td valign="top" align="left">Random Forest</td>
<td valign="top" align="left">Accuracy = 0.99&#x2003;AUC = 0.99 Sensitivity = 1.00 Specificity = 0.92</td>
</tr>
<tr>
<td valign="top" align="left">Sharif et&#xa0;al. (2020) (<xref ref-type="bibr" rid="B35">35</xref>)</td>
<td valign="top" align="left">HGG vs. LGG (unclear)</td>
<td valign="top" align="left">BraTS (n=30 patients)</td>
<td valign="top" align="left">2.3:1 (7 HGG, 3 LGG in testing set)</td>
<td valign="top" align="left">Internal (Holdout, 10-fold cross-validation)</td>
<td valign="top" align="left">T1, T1CE, T2, FLAIR</td>
<td valign="top" align="left">Deep learning extracted</td>
<td valign="top" align="left">Convolutional Neural Network</td>
<td valign="top" align="left">Accuracy = 0.987</td>
</tr>
<tr>
<td valign="top" align="left">Muneer et&#xa0;al. (2019) (<xref ref-type="bibr" rid="B36">36</xref>)</td>
<td valign="top" align="left">1 vs. 2 vs. 3 vs. 4</td>
<td valign="top" align="left">Single center hospital (n=20 patients)</td>
<td valign="top" align="left">1.3:1.6:1:1.5 (39 grade 1, 51 grade 2, 31 grade 3, 47 grade 4 images in testing set)</td>
<td valign="top" align="left">Internal<break/>(Holdout, 30% of dataset)</td>
<td valign="top" align="left">T2</td>
<td valign="top" align="left">Deep learning extracted</td>
<td valign="top" align="left">VGG19 (Deep Convolutional Neural Network)</td>
<td valign="top" align="left">Accuracy = 0.9825 Sensitivity = 0.9272 Specificity = 0.9813 Positive Predictive Value = 0.9471&#x2003;F1 Score = 0.9371</td>
</tr>
<tr>
<td valign="top" align="left">Dandil and Bicer (2020) (<xref ref-type="bibr" rid="B37">37</xref>)</td>
<td valign="top" align="left">1/2 vs. 3 vs. 4 vs. meningioma</td>
<td valign="top" align="left">INTERPRET (n=179 patients)</td>
<td valign="top" align="left">Unclear</td>
<td valign="top" align="left">Unspecified</td>
<td valign="top" align="left">MR Spectroscopy (Time of Echo 20ms and 136ms)</td>
<td valign="top" align="left">First-order, Shape and size, Texture</td>
<td valign="top" align="left">Long Short-Term Memory (Neural Network)</td>
<td valign="top" align="left">Accuracy = 0.982 AUC = 0.9936 Sensitivity = 1.00 Specificity = 0.9753</td>
</tr>
<tr>
<td valign="top" align="left">Tian et&#xa0;al. (2018) (<xref ref-type="bibr" rid="B38">38</xref>)</td>
<td valign="top" align="left">2 vs. 3/4</td>
<td valign="top" align="left">Single center hospital (n=153 patients)</td>
<td valign="top" align="left">2.6:1 (111 HGG, 42 LGG in total set)</td>
<td valign="top" align="left">Internal (10-fold cross-validation)</td>
<td valign="top" align="left">T1, T1CE, T2, Diffusion, Perfusion (3D Arterial Spin Labeling)</td>
<td valign="top" align="left">Texture</td>
<td valign="top" align="left">Support Vector Machine</td>
<td valign="top" align="left">Accuracy = 0.981 AUC = 0.992 Sensitivity = 0.987 Specificity = 0.974</td>
</tr>
<tr>
<td valign="top" align="left">Lo et&#xa0;al. (2019) (<xref ref-type="bibr" rid="B39">39</xref>)</td>
<td valign="top" align="left">2 vs. 3 vs. 4</td>
<td valign="top" align="left">TCIA (n=130 patients)</td>
<td valign="top" align="left">1:1.4:1.9(30 grade 2, 43 grade 3 and 57 grade 4 in total set)</td>
<td valign="top" align="left">Internal (10-fold cross-validation)</td>
<td valign="top" align="left">T1CE</td>
<td valign="top" align="left">Deep learning extracted</td>
<td valign="top" align="left">Deep Convolutional Neural Network</td>
<td valign="top" align="left">Accuracy = 0.979 AUC = 0.9991</td>
</tr>
<tr>
<td valign="top" align="left">Kumar et&#xa0;al. (2020) (<xref ref-type="bibr" rid="B40">40</xref>)</td>
<td valign="top" align="left">1/2 vs. 3/4</td>
<td valign="top" align="left">BraTS (n=285 patients)</td>
<td valign="top" align="left">2.8:1 (210 HGG, 75 LGG in total set)</td>
<td valign="top" align="left">Internal (5-fold cross-validation)</td>
<td valign="top" align="left">T1, T1CE, T2, (T2W)-FLAIR</td>
<td valign="top" align="left">First-order, Shape, Texture</td>
<td valign="top" align="left">Random Forest</td>
<td valign="top" align="left">Accuracy = 0.9754 AUC = 0.9748 Sensitivity = 0.9762 Specificity = 0.9733&#x2003;F1 Score = 0.983</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>Testing or validation metrics are reported when available, otherwise training metrics are reported. HGG, high-grade gliomas; LGG, low-grade gliomas; ML, machine learning; PRISMA-DTA, Preferred Reporting Items for Systematic Reviews and Meta-Analyses of Diagnostic Test Accuracy; T1CE, T1-weighted contrast-enhanced; TRIPOD, Transparent Reporting of a multivariable prediction model for Individual Prognosis Or Diagnosis.</p>
</fn>
</table-wrap-foot>
</table-wrap>
</sec>
</sec>
<sec id="s3_3">
<title>3.3 Quality Assessment</title>
<p>The mean adherence rate to TRIPOD was 44.5% &#xb1; 11.1%, with poor reporting adherence in categories including model performance (0%), abstract (0%), title (0%), justification of sample size (2.4%), full model specification (2.4%), and participant demographics and missing data (7.1%). High reporting adherence was observed for results interpretation (100%), background (98.8%), study design/source of data (96.5%), and objectives (95.3%) (<xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>).</p>
<fig id="f6" position="float">
<label>Figure&#xa0;6</label>
<caption>
<p>TRIPOD adherence of machine learning glioma grade prediction studies. Adherence rate for individual items represents the percent of studies scoring a point for that item: 1 &#x2013; title. 2 &#x2013; abstract. 3a &#x2013; background. 3b &#x2013; objectives. 4a &#x2013; study design. 4b &#x2013; study dates. 5a &#x2013; study setting. 5b &#x2013; eligibility criteria. 6a &#x2013; outcome assessment. 6b &#x2013; blinding assessment of outcome. 7a &#x2013; predictor assessment. 7b &#x2013; blinding assessment of predictors. 8 &#x2013; sample size justification. 9 &#x2013; missing data. 10a &#x2013; predictor handling. 10b &#x2013; model type, model-building, and internal validation. 10d &#x2013; model performance. 13a &#x2013; participant flow and outcomes. 13b &#x2013; participant demographics and missing data. 14a &#x2013; model development (participants and outcomes). 15a &#x2013; full model specification. 15b &#x2013; using the model. 16 &#x2013; model performance. 18 &#x2013; study limitations. 19b &#x2013; results interpretation. 20 &#x2013; clinical use and research implications. 22 &#x2013; funding. Overall &#x2013; mean TRIPOD adherence rate of all studies. TRIPOD, Transparent Reporting of a multivariable prediction model for Individual Prognosis Or Diagnosis.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fonc-12-856231-g006.tif"/>
</fig>
<p>Titles (0%) did not identify the development of prediction models. Abstracts (0%) frequently lacked source of data, overall sample size, and calibration methods. Regarding model performance (0%), very few studies reported measures for model calibration or confidence intervals. Most studies failed to specify model regression coefficients (2.4%) or provide justifications of sample size (2.4%), e.g., how sample size was arrived at according to statistical or practical grounds. More detailed explanations of TRIPOD items and their adherence rates can be found in <xref ref-type="supplementary-material" rid="ST1">
<bold>Table A8 (Supplementary)</bold>
</xref>.</p>
</sec>
</sec>
<sec id="s4">
<title>4 Discussion</title>
<p>Our systematic review analyzed 85 articles describing ML applications for glioma grade prediction and revealed several trends. First, the number of studies published per year grew steadily between 2016 and 2019. Second, imaging sequences and ML models became less conventional, with the emergence of advanced MRI sequences (MR Spectroscopy, Perfusion) in the early 2000s (<xref ref-type="bibr" rid="B41">41</xref>) and DL models in 2018 (<xref ref-type="bibr" rid="B42">42</xref>). Third, datasets recently expanded to encompass multiple institutions, with BraTS and TCIA datasets appearing in 2017. While ML model studies report high predictive accuracies, they underreport critical model performance measures, lack a common validation dataset, and vary remarkably in glioma classification systems, ML algorithms, feature types and imaging sequences used for prediction, etc., limiting model comparison. Here, we identify several opportunities for improvement to prepare models for multicenter clinical adoption.</p>
<sec id="s4_1">
<title>4.1 Study Datasets, Validation Techniques, Classification Systems, and Reporting Quality</title>
<p>Prior to broad clinical use, ML models must be trained and validated on large, multi-institutional datasets to ensure generalizability (<xref ref-type="bibr" rid="B43">43</xref>). Dataset sizes, however, were low in our study, and most publications lacked external validation. These findings are consistent with those from a similar systematic review by Tabatabaei et&#xa0;al. (<xref ref-type="bibr" rid="B23">23</xref>). Moreover, while studies based on highly curated datasets, including BraTS or TCIA, showed consistently high accuracy results, algorithms trained on these datasets without external validation may not have reproducible results in clinical practice, where imaging protocols are less standardized, image quality is variable, and tumor presentations are heterogeneous. To show that models perform well across distinct populations and are fit for broad clinical implementation, future works should use sizable, less-curated, multicenter datasets and externally validate their models.</p>
<p>ML models should also be trained according to standardized definitions of glioma grade. Interestingly, definitions were variable for HGG and LGG, with some studies considering grade 3 gliomas to be high-grade and others low-grade. Lack of a unified classification system may hinder predictive model performances on external datasets, given that the images used for segmentation, feature extraction, and model training/testing are labeled HGG or LGG based on non-uniform definitions. As glioma grade guides clinical management, it is essential that algorithm outputs of &#x201c;HGG&#x201d; and &#x201c;LGG&#x201d; reflect a universal definition consistent with current WHO criteria.</p>
<p>An alternative to binary high-or-low grading is to report numerical glioma grades (1, 2, 3, or 4) and tumor entities. Importantly, grade and entity classifications are evolving. In 2016, purely histopathological classification was succeeded by classification based on both molecular and histopathologic parameters (<xref ref-type="bibr" rid="B44">44</xref>). In 2021, cIMPACT-NOW established further changes to glioma grading, for example, by redefining GBM to be an IDH (Isocitrate Dehydrogenase)-wild-type lesion distinct from IDH-mutant grade 4 astrocytomas (<xref ref-type="bibr" rid="B45">45</xref>, <xref ref-type="bibr" rid="B46">46</xref>). Classification changes may have led to inconsistencies in tumor entities and grades reported in glioma grade prediction studies across the years, limiting model comparison. As WHO criteria continue to evolve and affect generalizability of study results, we recommend that studies clearly reference the criteria used in glioma grading, report the glioma entities and corresponding grade used in model development, and describe the predictive performances by both entity and grade. This will promote comprehensible and traceable results over time. Moreover, integrated techniques that characterize disease according to both radiological and biological features are emerging in neuro-oncology (<xref ref-type="bibr" rid="B47">47</xref>). We advise future researchers to consider the implementation of these techniques into ML model development studies predicting glioma grade, molecular markers, response to treatment, prognosis, and other applications within neuro-oncology.</p>
<p>Finally, reporting of ML models should be transparent, thorough, and reproducible to facilitate proper assessment for use in clinical practice (<xref ref-type="bibr" rid="B48">48</xref>). Several comprehensive checklists are used to assess reporting quality of diagnostic models, including Checklist for Artificial Intelligence in Medical Imaging (<xref ref-type="bibr" rid="B49">49</xref>) and TRIPOD. In our study, mean adherence to TRIPOD was low, with key assessment elements such as model performance scoring poorly. These findings reflect inadequate study reporting. To address this, we recommend future studies use appropriate reporting frameworks to guide all phases of study execution, from initial design through manuscript writing. For ML studies, the relevance of TRIPOD as a benchmark for reporting quality may be questioned. Published explanations and elaborations of TRIPOD focus on regression-based models, a shortcoming that TRIPOD authors have recently acknowledged (<xref ref-type="bibr" rid="B19">19</xref>). We support their initiative to create a TRIPOD Statement specific to ML (TRIPOD-AI) (<xref ref-type="bibr" rid="B19">19</xref>, <xref ref-type="bibr" rid="B20">20</xref>), and in the context of this work, to improve the reporting quality of literature concerning ML in glioma grade prediction.</p>
</sec>
<sec id="s4_2">
<title>4.2 Limitations</title>
<p>This study has several limitations. First, the timing and criteria of our search may have missed relevant studies (e.g., recent and unpublished works). Moreover, 191 of the 1,135 (16.8%) studies assessed for eligibility were excluded because they were abstracts (n=169, 14.9%) or not in English (n=22, 1.9%), creating a potential selection bias. However, full texts were required for complete data extraction and quality of reporting analysis, and we unfortunately did not have the resources to translate non-English articles. Second, we determined best performing algorithms based on accuracy, which excluded the three studies that did not report accuracy results for their models. Accuracy, furthermore, may be considered a flawed performance metric for ML models applied to imbalanced datasets (<xref ref-type="bibr" rid="B50">50</xref>), which constituted most datasets in our study. With imbalanced datasets, ML models intrinsically overfit toward the majority class, risking higher misclassification rates for minority classes (<xref ref-type="bibr" rid="B51">51</xref>, <xref ref-type="bibr" rid="B52">52</xref>). Because accuracy may be high even if a minority class is poorly predicted, we recommend study authors consistently report a full slate of model performance metrics. Including metrics sensitive to performance differences within imbalanced datasets (e.g., AUC) (<xref ref-type="bibr" rid="B53">53</xref>) will enable a more thorough assessment of ML model performance. Third, the inconsistent definitions for HGG and LGG, evolving grading criteria, high heterogeneity of our included articles and low number of articles reporting confidence intervals for their performance metrics limited the pooling of results across studies and subsequent generation of conclusions. As a result, we could not perform a meta-analysis (<xref ref-type="bibr" rid="B54">54</xref>, <xref ref-type="bibr" rid="B55">55</xref>).</p>
</sec>
</sec>
<sec id="s5">
<title>5 Conclusion</title>
<p>The application of ML to glioma grade prediction has grown substantially, with ML model studies reporting high predictive accuracies but lacking essential metrics and characteristics for assessing model performance. To increase the generalizability, standardization, reproducibility, and reporting quality necessary for clinical translation, future studies need to (1) train and test on large, multi-institutional datasets, (2) validate on external datasets, (3) clearly report glioma entities, corresponding glioma grades, and a full state of predictive performance metrics by both grade and entity, and (4) adhere to reporting guidelines.</p>
</sec>
<sec sec-type="data-availability" id="s6">
<title>Data Availability Statement</title>
<p>The original contributions presented in the study are included in the article/<xref ref-type="supplementary-material" rid="ST1">
<bold>Supplementary Material</bold>
</xref>, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec sec-type="author-contributions" id="s7">
<title>Author Contributions</title>
<p>All authors listed have made a substantial, direct, and intellectual contribution to the work and approved it for publication.</p>
</sec>
<sec sec-type="funding-information" id="s8">
<title>Funding</title>
<p>SM receives funding in part from the Biomedical Education Program (BMEP). RB receives funding in part from the National Institute of Diabetes and Digestive and Kidney Disease of the National Institutes of Health under Award Number T35DK104689. MA received funding from American Society of Neuroradiology Fellow Award 2018. This publication was made possible by KL2 TR001862 (MA) from the National Center for Advancing Translational Science (NCATS), components of the National Institutes of Health (NIH), and NIH roadmap for Medical Research. Seyedmehdi Payabvash has grant support from NIH/NINDS K23NS118056, foundation of American Society of Neuroradiology (ASNR) #1861150721, Doris Duke Charitable Foundation (DDCF) #2020097, and NVIDIA. Funders were not involved in the study design, collection, analysis, interpretation of data, the writing of this article or the decision to submit it for publication.</p>
</sec>
<sec id="s9">
<title>Author Disclaimer</title>
<p>The content is solely the responsibility of the authors and does not necessarily represent the official views of the National Institutes of Health.</p>
</sec>
<sec id="s10" sec-type="COI-statement">
<title>Conflict of Interest</title>
<p>Author ML is an employee and stockholder of Visage Imaging, Inc., and unrelated to this work, receives funding from NIH/NCI R01 CA206180 and is a board member of Tau Beta Pi engineering honor society. KB is an employee of Visage Imaging, GmbH. JI has funding support for an investigator-initiated clinical trial from Novartis Pharmaceuticals (unrelated to this work).</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s11" sec-type="disclaimer">
<title>Publisher&#x2019;s Note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
</body>
<back>
<ack>
<title>Acknowledgments</title>
<p>We would like to thank our institutional librarians (Alexandria Brackett and Thomas Mead) for developing and executing our search strategy, and Mary Hughes and Vermetha Polite for their technical support. We also would like to thank Julia Shatalov for assisting with eligibility screening and Irena Tocino for her continuous support throughout the execution of this study.</p>
</ack>
<sec id="s12" sec-type="supplementary-material">
<title>Supplementary Material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fonc.2022.856231/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fonc.2022.856231/full#supplementary-material</ext-link>
</p>
  <supplementary-material xlink:href="Table_1.docx" id="ST1" mimetype="application/vnd.openxmlformats-officedocument.wordprocessingml.document"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<label>1</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ostrom</surname> <given-names>QT</given-names>
</name>
<name>
<surname>Cioffi</surname> <given-names>G</given-names>
</name>
<name>
<surname>Gittleman</surname> <given-names>H</given-names>
</name>
<name>
<surname>Patil</surname> <given-names>N</given-names>
</name>
<name>
<surname>Waite</surname> <given-names>K</given-names>
</name>
<name>
<surname>Kruchko</surname> <given-names>C</given-names>
</name>
<etal/>
</person-group>. <article-title>CBTRUS Statistical Report: Primary Brain and Other Central Nervous System Tumors Diagnosed in the United States in 2012-2016</article-title>. <source>Neuro Oncol</source> (<year>2019</year>) <volume>21</volume>(<supplement>Suppl 5</supplement>):<fpage>v1</fpage>&#x2013;<lpage>v100</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/neuonc/noz150</pub-id>
</citation>
</ref>
<ref id="B2">
<label>2</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Louis</surname> <given-names>DN</given-names>
</name>
<name>
<surname>Perry</surname> <given-names>A</given-names>
</name>
<name>
<surname>Wesseling</surname> <given-names>P</given-names>
</name>
<name>
<surname>Brat</surname> <given-names>DJ</given-names>
</name>
<name>
<surname>Cree</surname> <given-names>IA</given-names>
</name>
<name>
<surname>Figarella-Branger</surname> <given-names>D</given-names>
</name>
<etal/>
</person-group>. <article-title>The 2021 WHO Classification of Tumors of the Central Nervous System: A Summary</article-title>. <source>Neuro Oncol</source> (<year>2021</year>) <volume>23</volume>(<issue>8</issue>):<page-range>1231&#x2013;51</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/neuonc/noab106</pub-id>
</citation>
</ref>
<ref id="B3">
<label>3</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tran</surname> <given-names>B</given-names>
</name>
<name>
<surname>Rosenthal</surname> <given-names>MA</given-names>
</name>
</person-group>. <article-title>Survival Comparison Between Glioblastoma Multiforme and Other Incurable Cancers</article-title>. <source>J Clin Neurosci</source> (<year>2010</year>) <volume>17</volume>(<issue>4</issue>):<page-range>417&#x2013;21</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.jocn.2009.09.004</pub-id>
</citation>
</ref>
<ref id="B4">
<label>4</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ohgaki</surname> <given-names>H</given-names>
</name>
<name>
<surname>Kleihues</surname> <given-names>P</given-names>
</name>
</person-group>. <article-title>Population-Based Studies on Incidence, Survival Rates, and Genetic Alterations in Astrocytic and Oligodendroglial Gliomas</article-title>. <source>J Neuropathol Exp Neurol</source> (<year>2005</year>) <volume>64</volume>(<issue>6</issue>):<page-range>479&#x2013;89</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/jnen/64.6.479</pub-id>
</citation>
</ref>
<ref id="B5">
<label>5</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gallego Perez-Larraya</surname> <given-names>J</given-names>
</name>
<name>
<surname>Delattre</surname> <given-names>JY</given-names>
</name>
</person-group>. <article-title>Management of Elderly Patients With Gliomas</article-title>. <source>Oncologist</source> (<year>2014</year>) <volume>19</volume>(<issue>12</issue>):<page-range>1258&#x2013;67</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1634/theoncologist.2014-0170</pub-id>
</citation>
</ref>
<ref id="B6">
<label>6</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zonari</surname> <given-names>P</given-names>
</name>
<name>
<surname>Baraldi</surname> <given-names>P</given-names>
</name>
<name>
<surname>Crisi</surname> <given-names>G</given-names>
</name>
</person-group>. <article-title>Multimodal MRI in the Characterization of Glial Neoplasms: The Combined Role of Single-Voxel MR Spectroscopy, Diffusion Imaging and Echo-Planar Perfusion Imaging</article-title>. <source>Neuroradiology</source> (<year>2007</year>) <volume>49</volume>(<issue>10</issue>):<fpage>795</fpage>&#x2013;<lpage>803</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s00234-007-0253-x</pub-id>
</citation>
</ref>
<ref id="B7">
<label>7</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Thon</surname> <given-names>N</given-names>
</name>
<name>
<surname>Tonn</surname> <given-names>JC</given-names>
</name>
<name>
<surname>Kreth</surname> <given-names>FW</given-names>
</name>
</person-group>. <article-title>The Surgical Perspective in Precision Treatment of Diffuse Gliomas</article-title>. <source>Onco Targets Ther</source> (<year>2019</year>) <volume>12</volume>:<page-range>1497&#x2013;508</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.2147/OTT.S174316</pub-id>
</citation>
</ref>
<ref id="B8">
<label>8</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hu</surname> <given-names>LS</given-names>
</name>
<name>
<surname>Hawkins-Daarud</surname> <given-names>A</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>L</given-names>
</name>
<name>
<surname>Li</surname> <given-names>J</given-names>
</name>
<name>
<surname>Swanson</surname> <given-names>KR</given-names>
</name>
</person-group>. <article-title>Imaging of Intratumoral Heterogeneity in High-Grade Glioma</article-title>. <source>Cancer Lett</source> (<year>2020</year>) <volume>477</volume>:<fpage>97</fpage>&#x2013;<lpage>106</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.canlet.2020.02.025</pub-id>
</citation>
</ref>
<ref id="B9">
<label>9</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Law</surname> <given-names>M</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>S</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>H</given-names>
</name>
<name>
<surname>Babb</surname> <given-names>JS</given-names>
</name>
<name>
<surname>Johnson</surname> <given-names>G</given-names>
</name>
<name>
<surname>Cha</surname> <given-names>S</given-names>
</name>
<etal/>
</person-group>. <article-title>Glioma Grading: Sensitivity, Specificity, and Predictive Values of Perfusion MR Imaging and Proton MR Spectroscopic Imaging Compared With Conventional MR Imaging</article-title>. <source>AJNR Am J Neuroradiol</source> (<year>2003</year>) <volume>24</volume>(<issue>10</issue>):<page-range>1989&#x2013;98</page-range>.</citation>
</ref>
<ref id="B10">
<label>10</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Gillies</surname> <given-names>RJ</given-names>
</name>
<name>
<surname>Kinahan</surname> <given-names>PE</given-names>
</name>
<name>
<surname>Hricak</surname> <given-names>H</given-names>
</name>
</person-group>. <article-title>Radiomics: Images Are More Than Pictures, They Are Data</article-title>. <source>Radiology</source> (<year>2016</year>) <volume>278</volume>(<issue>2</issue>):<page-range>563&#x2013;77</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1148/radiol.2015151169</pub-id>
</citation>
</ref>
<ref id="B11">
<label>11</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Giger</surname> <given-names>ML</given-names>
</name>
</person-group>. <article-title>Machine Learning in Medical Imaging</article-title>. <source>J Am Coll Radiol</source> (<year>2018</year>) <volume>15</volume>(<issue>3 Pt B</issue>):<page-range>512&#x2013;20</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.jacr.2017.12.028</pub-id>
</citation>
</ref>
<ref id="B12">
<label>12</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chartrand</surname> <given-names>G</given-names>
</name>
<name>
<surname>Cheng</surname> <given-names>PM</given-names>
</name>
<name>
<surname>Vorontsov</surname> <given-names>E</given-names>
</name>
<name>
<surname>Drozdzal</surname> <given-names>M</given-names>
</name>
<name>
<surname>Turcotte</surname> <given-names>S</given-names>
</name>
<name>
<surname>Pal</surname> <given-names>CJ</given-names>
</name>
<etal/>
</person-group>. <article-title>Deep Learning: A Primer for Radiologists</article-title>. <source>Radiographics</source> (<year>2017</year>) <volume>37</volume>(<issue>7</issue>):<page-range>2113&#x2013;31</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1148/rg.2017170077</pub-id>
</citation>
</ref>
<ref id="B13">
<label>13</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname> <given-names>PM</given-names>
</name>
<name>
<surname>Montagnon</surname> <given-names>E</given-names>
</name>
<name>
<surname>Yamashita</surname> <given-names>R</given-names>
</name>
<name>
<surname>Pan</surname> <given-names>I</given-names>
</name>
<name>
<surname>Cadrin-Ch&#xea;nevert</surname> <given-names>A</given-names>
</name>
<name>
<surname>Perdig&#xf3;n Romero</surname> <given-names>F</given-names>
</name>
<etal/>
</person-group>. <article-title>Deep Learning: An Update for Radiologists</article-title>. <source>Radiographics</source> (<year>2021</year>) <volume>41</volume>(<issue>5</issue>):<page-range>1427&#x2013;45</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1148/rg.2021200210</pub-id>
</citation>
</ref>
<ref id="B14">
<label>14</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>He</surname> <given-names>J</given-names>
</name>
<name>
<surname>Baxter</surname> <given-names>SL</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>J</given-names>
</name>
<name>
<surname>Xu</surname> <given-names>J</given-names>
</name>
<name>
<surname>Zhou</surname> <given-names>X</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>K</given-names>
</name>
</person-group>. <article-title>The Practical Implementation of Artificial Intelligence Technologies in Medicine</article-title>. <source>Nat Med</source> (<year>2019</year>) <volume>25</volume>(<issue>1</issue>):<page-range>30&#x2013;6</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41591-018-0307-0</pub-id>
</citation>
</ref>
<ref id="B15">
<label>15</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lasocki</surname> <given-names>A</given-names>
</name>
<name>
<surname>Tsui</surname> <given-names>A</given-names>
</name>
<name>
<surname>Tacey</surname> <given-names>MA</given-names>
</name>
<name>
<surname>Drummond</surname> <given-names>KJ</given-names>
</name>
<name>
<surname>Field</surname> <given-names>KM</given-names>
</name>
<name>
<surname>Gaillard</surname> <given-names>F</given-names>
</name>
</person-group>. <article-title>MRI Grading Versus Histology: Predicting Survival of World Health Organization Grade II-IV Astrocytomas</article-title>. <source>AJNR Am J Neuroradiol</source> (<year>2015</year>) <volume>36</volume>(<issue>1</issue>):<fpage>77</fpage>&#x2013;<lpage>83</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3174/ajnr.A4077</pub-id>
</citation>
</ref>
<ref id="B16">
<label>16</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jin</surname> <given-names>W</given-names>
</name>
<name>
<surname>Fatehi</surname> <given-names>M</given-names>
</name>
<name>
<surname>Abhishek</surname> <given-names>K</given-names>
</name>
<name>
<surname>Mallya</surname> <given-names>M</given-names>
</name>
<name>
<surname>Toyota</surname> <given-names>B</given-names>
</name>
<name>
<surname>Hamarneh</surname> <given-names>G</given-names>
</name>
</person-group>. <article-title>Artificial Intelligence in Glioma Imaging: Challenges and Advances</article-title>. <source>J Neural Eng</source> (<year>2020</year>) <volume>17</volume>(<issue>2</issue>):<elocation-id>021002</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1088/1741-2552/ab8131</pub-id>
</citation>
</ref>
<ref id="B17">
<label>17</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Park</surname> <given-names>JE</given-names>
</name>
<name>
<surname>Kim</surname> <given-names>HS</given-names>
</name>
<name>
<surname>Kim</surname> <given-names>D</given-names>
</name>
<name>
<surname>Park</surname> <given-names>SY</given-names>
</name>
<name>
<surname>Kim</surname> <given-names>JY</given-names>
</name>
<name>
<surname>Cho</surname> <given-names>SJ</given-names>
</name>
<etal/>
</person-group>. <article-title>A Systematic Review Reporting Quality of Radiomics Research in Neuro-Oncology: Toward Clinical Utility and Quality Improvement Using High-Dimensional Imaging Features</article-title>. <source>BMC Cancer</source> (<year>2020</year>) <volume>20</volume>(<issue>1</issue>):<fpage>29</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1186/s12885-019-6504-5</pub-id>
</citation>
</ref>
<ref id="B18">
<label>18</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Collins</surname> <given-names>GS</given-names>
</name>
<name>
<surname>Reitsma</surname> <given-names>JB</given-names>
</name>
<name>
<surname>Altman</surname> <given-names>DG</given-names>
</name>
<name>
<surname>Moons</surname> <given-names>KG</given-names>
</name>
</person-group>. <article-title>Transparent Reporting of a Multivariable Prediction Model for Individual Prognosis or Diagnosis (TRIPOD): The TRIPOD Statement</article-title>. <source>Ann Intern Med</source> (<year>2015</year>) <volume>162</volume>(<issue>1</issue>):<fpage>55</fpage>&#x2013;<lpage>63</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.7326/M14-0697</pub-id>
</citation>
</ref>
<ref id="B19">
<label>19</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Collins</surname> <given-names>GS</given-names>
</name>
<name>
<surname>Moons</surname> <given-names>KGM</given-names>
</name>
</person-group>. <article-title>Reporting of Artificial Intelligence Prediction Models</article-title>. <source>Lancet</source> (<year>2019</year>) <volume>393</volume>(<issue>10181</issue>):<page-range>1577&#x2013;9</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/S0140-6736(19)30037-6</pub-id>
</citation>
</ref>
<ref id="B20">
<label>20</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Collins</surname> <given-names>GS</given-names>
</name>
<name>
<surname>Dhiman</surname> <given-names>P</given-names>
</name>
<name>
<surname>Andaur Navarro</surname> <given-names>CL</given-names>
</name>
<name>
<surname>Ma</surname> <given-names>J</given-names>
</name>
<name>
<surname>Hooft</surname> <given-names>L</given-names>
</name>
<name>
<surname>Reitsma</surname> <given-names>JB</given-names>
</name>
<etal/>
</person-group>. <article-title>Protocol for Development of a Reporting Guideline (TRIPOD-AI) and Risk of Bias Tool (PROBAST-AI) for Diagnostic and Prognostic Prediction Model Studies Based on Artificial Intelligence</article-title>. <source>BMJ Open</source> (<year>2021</year>) <volume>11</volume>(<issue>7</issue>):<elocation-id>e048008</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1136/bmjopen-2020-048008</pub-id>
</citation>
</ref>
<ref id="B21">
<label>21</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Buchlak</surname> <given-names>QD</given-names>
</name>
<name>
<surname>Esmaili</surname> <given-names>N</given-names>
</name>
<name>
<surname>Leveque</surname> <given-names>JC</given-names>
</name>
<name>
<surname>Bennett</surname> <given-names>C</given-names>
</name>
<name>
<surname>Farrokhi</surname> <given-names>F</given-names>
</name>
<name>
<surname>Piccardi</surname> <given-names>M</given-names>
</name>
</person-group>. <article-title>Machine Learning Applications to Neuroimaging for Glioma Detection and Classification: An Artificial Intelligence Augmented Systematic Review</article-title>. <source>J Clin Neurosci</source> (<year>2021</year>) <volume>89</volume>:<page-range>177&#x2013;98</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.jocn.2021.04.043</pub-id>
</citation>
</ref>
<ref id="B22">
<label>22</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sohn</surname> <given-names>CK</given-names>
</name>
<name>
<surname>Bisdas</surname> <given-names>S</given-names>
</name>
</person-group>. <article-title>Diagnostic Accuracy of Machine Learning-Based Radiomics in Grading Gliomas: Systematic Review and Meta-Analysis</article-title>. <source>Contrast Media Mol Imaging</source> (<year>2020</year>) <volume>2020</volume>:<elocation-id>2127062</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1155/2020/2127062</pub-id>
</citation>
</ref>
<ref id="B23">
<label>23</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tabatabaei</surname> <given-names>M</given-names>
</name>
<name>
<surname>Razaei</surname> <given-names>A</given-names>
</name>
<name>
<surname>Sarrami</surname> <given-names>AH</given-names>
</name>
<name>
<surname>Saadatpour</surname> <given-names>Z</given-names>
</name>
<name>
<surname>Singhal</surname> <given-names>A</given-names>
</name>
<name>
<surname>Sotoudeh</surname> <given-names>H</given-names>
</name>
</person-group>. <article-title>Current Status and Quality of Machine Learning-Based Radiomics Studies for Glioma Grading: A Systematic Review</article-title>. <source>Oncology</source> (<year>2021</year>) <volume>99</volume>(<issue>7</issue>):<fpage>433</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1159/000515597</pub-id>
</citation>
</ref>
<ref id="B24">
<label>24</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Frank</surname> <given-names>RA</given-names>
</name>
<name>
<surname>Bossuyt</surname> <given-names>PM</given-names>
</name>
<name>
<surname>McInnes</surname> <given-names>MDF</given-names>
</name>
</person-group>. <article-title>Systematic Reviews and Meta-Analyses of Diagnostic Test Accuracy: The PRISMA-DTA Statement</article-title>. <source>Radiology</source> (<year>2018</year>) <volume>289</volume>(<issue>2</issue>):<page-range>313&#x2013;4</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1148/radiol.2018180850</pub-id>
</citation>
</ref>
<ref id="B25">
<label>25</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Whiting</surname> <given-names>P</given-names>
</name>
<name>
<surname>Westwood</surname> <given-names>M</given-names>
</name>
<name>
<surname>Burke</surname> <given-names>M</given-names>
</name>
<name>
<surname>Sterne</surname> <given-names>J</given-names>
</name>
<name>
<surname>Glanville</surname> <given-names>J</given-names>
</name>
</person-group>. <article-title>Systematic Reviews of Test Accuracy Should Search a Range of Databases to Identify Primary Studies</article-title>. <source>J Clin Epidemiol</source> (<year>2008</year>) <volume>61</volume>(<issue>4</issue>):<page-range>357&#x2013;64</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.jclinepi.2007.05.013</pub-id>
</citation>
</ref>
<ref id="B26">
<label>26</label>
<citation citation-type="web">
<person-group person-group-type="author">
<collab>TRIPOD Statement Web Site</collab>
</person-group>. <source>Adherence to TRIPOD</source> (<year>2020</year>). Available at: <uri xlink:href="https://www.tripod-statement.org/adherence/">https://www.tripod-statement.org/adherence/</uri> (Accessed <access-date>10 Jul 2021</access-date>).</citation>
</ref>
<ref id="B27">
<label>27</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Moons</surname> <given-names>KG</given-names>
</name>
<name>
<surname>Altman</surname> <given-names>DG</given-names>
</name>
<name>
<surname>Reitsma</surname> <given-names>JB</given-names>
</name>
<name>
<surname>Ioannidis</surname> <given-names>JP</given-names>
</name>
<name>
<surname>Macaskill</surname> <given-names>P</given-names>
</name>
<name>
<surname>Steyerberg</surname> <given-names>EW</given-names>
</name>
<etal/>
</person-group>. <article-title>Transparent Reporting of a Multivariable Prediction Model for Individual Prognosis or Diagnosis (TRIPOD): Explanation and Elaboration</article-title>. <source>Ann Intern Med</source> (<year>2015</year>) <volume>162</volume>(<issue>1</issue>):<fpage>W1</fpage>&#x2013;<lpage>73</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.7326/M14-0698</pub-id>
</citation>
</ref>
<ref id="B28">
<label>28</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Menze</surname> <given-names>BH</given-names>
</name>
<name>
<surname>Jakab</surname> <given-names>A</given-names>
</name>
<name>
<surname>Bauer</surname> <given-names>S</given-names>
</name>
<name>
<surname>Kalpathy-Cramer</surname> <given-names>J</given-names>
</name>
<name>
<surname>Farahani</surname> <given-names>K</given-names>
</name>
<name>
<surname>Kirby</surname> <given-names>J</given-names>
</name>
<etal/>
</person-group>. <article-title>The Multimodal Brain Tumor Image Segmentation Benchmark (BRATS)</article-title>. <source>IEEE T Med Imaging</source> (<year>2015</year>) <volume>34</volume>(<issue>10</issue>):<fpage>1993</fpage>&#x2013;<lpage>2024</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/Tmi.2014.2377694</pub-id>
</citation>
</ref>
<ref id="B29">
<label>29</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Clark</surname> <given-names>K</given-names>
</name>
<name>
<surname>Vendt</surname> <given-names>B</given-names>
</name>
<name>
<surname>Smith</surname> <given-names>K</given-names>
</name>
<name>
<surname>Freymann</surname> <given-names>J</given-names>
</name>
<name>
<surname>Kirby</surname> <given-names>J</given-names>
</name>
<name>
<surname>Koppel</surname> <given-names>P</given-names>
</name>
<etal/>
</person-group>. <article-title>The Cancer Imaging Archive (TCIA): Maintaining and Operating a Public Information Repository</article-title>. <source>J Digital Imaging</source> (<year>2013</year>) <volume>26</volume>(<issue>6</issue>):<page-range>1045&#x2013;57</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s10278-013-9622-7</pub-id>
</citation>
</ref>
<ref id="B30">
<label>30</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Park</surname> <given-names>YW</given-names>
</name>
<name>
<surname>Choi</surname> <given-names>YS</given-names>
</name>
<name>
<surname>Ahn</surname> <given-names>SS</given-names>
</name>
<name>
<surname>Chang</surname> <given-names>JH</given-names>
</name>
<name>
<surname>Kim</surname> <given-names>SH</given-names>
</name>
<name>
<surname>Lee</surname> <given-names>SK</given-names>
</name>
</person-group>. <article-title>Radiomics MRI Phenotyping With Machine Learning to Predict the Grade of Lower-Grade Gliomas: A Study Focused on Nonenhancing Tumors</article-title>. <source>Korean J Radiol</source> (<year>2019</year>) <volume>20</volume>(<issue>9</issue>):<page-range>1381&#x2013;9</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.3348/kjr.2018.0814</pub-id>
</citation>
</ref>
<ref id="B31">
<label>31</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hedyehzadeh</surname> <given-names>M</given-names>
</name>
<name>
<surname>Nezhad</surname> <given-names>SYD</given-names>
</name>
<name>
<surname>Safdarian</surname> <given-names>N</given-names>
</name>
</person-group>. <article-title>Evaluation of Conventional Machine Learning Methods for Brain Tumour Type Classification</article-title>. <source>Cr Acad Bulg Sci</source> (<year>2020</year>) <volume>73</volume>(<issue>6</issue>):<page-range>856&#x2013;65</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.7546/Crabs.2020.06.14</pub-id>
</citation>
</ref>
<ref id="B32">
<label>32</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bashir Gonbadi</surname> <given-names>F</given-names>
</name>
<name>
<surname>Khotanlou</surname> <given-names>H</given-names>
</name>
</person-group>. <article-title>Glioma Brain Tumors Diagnosis and Classification in MR Images Based on Convolutional Neural Networks</article-title>. <conf-name>9th International Conference on Computer and Knowledge Engineering (Iccke 2019)</conf-name> (<year>2019</year>) <fpage>1</fpage>&#x2013;<lpage>5</lpage>. doi: <pub-id pub-id-type="doi">10.1109/ICCKE48569.2019.8965143</pub-id>
</citation>
</ref>
<ref id="B33">
<label>33</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Polly</surname> <given-names>FP</given-names>
</name>
<name>
<surname>Shil</surname> <given-names>SK</given-names>
</name>
<name>
<surname>Hossain</surname> <given-names>MA</given-names>
</name>
<name>
<surname>Ayman</surname> <given-names>A</given-names>
</name>
<name>
<surname>Jang</surname> <given-names>YM</given-names>
</name>
</person-group>. <article-title>Detection and Classification of HGG and LGG Brain Tumor Using Machine Learning</article-title>. <conf-name>32nd International Conference on Information Networking (Icoin)</conf-name>, (<year>2018</year>) <page-range>813&#x2013;7</page-range>. doi: <pub-id pub-id-type="doi">10.1109/ICOIN.2018.8343231</pub-id>
</citation>
</ref>
<ref id="B34">
<label>34</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>De Looze</surname> <given-names>C</given-names>
</name>
<name>
<surname>Beausang</surname> <given-names>A</given-names>
</name>
<name>
<surname>Cryan</surname> <given-names>J</given-names>
</name>
<name>
<surname>Loftus</surname> <given-names>T</given-names>
</name>
<name>
<surname>Buckley</surname> <given-names>PG</given-names>
</name>
<name>
<surname>Farrell</surname> <given-names>M</given-names>
</name>
<etal/>
</person-group>. <article-title>Machine Learning: A Useful Radiological Adjunct in Determination of a Newly Diagnosed Glioma's Grade and IDH Status</article-title>. <source>J Neuro-Oncol</source> (<year>2018</year>) <volume>139</volume>(<issue>2</issue>):<page-range>491&#x2013;9</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s11060-018-2895-4</pub-id>
</citation>
</ref>
<ref id="B35">
<label>35</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sharif</surname> <given-names>MI</given-names>
</name>
<name>
<surname>Li</surname> <given-names>JP</given-names>
</name>
<name>
<surname>Khan</surname> <given-names>MA</given-names>
</name>
<name>
<surname>Saleem</surname> <given-names>MA</given-names>
</name>
</person-group>. <article-title>Active Deep Neural Network Features Selection for Segmentation and Recognition of Brain Tumors Using MRI Images</article-title>. <source>Pattern Recogn Lett</source> (<year>2020</year>) <volume>129</volume>:<page-range>181&#x2013;9</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.patrec.2019.11.019</pub-id>
</citation>
</ref>
<ref id="B36">
<label>36</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Muneer</surname> <given-names>KVA</given-names>
</name>
<name>
<surname>Rajendran</surname> <given-names>VR</given-names>
</name>
<name>
<surname>Joseph</surname> <given-names>KP</given-names>
</name>
</person-group>. <article-title>Glioma Tumor Grade Identification Using Artificial Intelligent Techniques</article-title>. <source>J Med Syst</source> (<year>2019</year>) <volume>43</volume>(<issue>5</issue>). doi:&#xa0;ARTN 113 10.1007/s10916-019-1228-2
</citation>
</ref>
<ref id="B37">
<label>37</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dandil</surname> <given-names>E</given-names>
</name>
<name>
<surname>Bicer</surname> <given-names>A</given-names>
</name>
</person-group>. <article-title>Automatic Grading of Brain Tumours Using LSTM Neural Networks on Magnetic Resonance Spectroscopy Signals</article-title>. <source>Iet Image Process</source>(<year>2020</year>) <volume>14</volume>(<issue>10</issue>):<page-range>1967&#x2013;79</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1049/iet-ipr.2019.1416</pub-id>
</citation>
</ref>
<ref id="B38">
<label>38</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tian</surname> <given-names>Q</given-names>
</name>
<name>
<surname>Yan</surname> <given-names>LF</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>X</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>X</given-names>
</name>
<name>
<surname>Hu</surname> <given-names>YC</given-names>
</name>
<name>
<surname>Han</surname> <given-names>Y</given-names>
</name>
<etal/>
</person-group>. <article-title>Radiomics Strategy for Glioma Grading Using Texture Features From Multiparametric MRI</article-title>. <source>J Magnetic Resonance Imaging</source> (<year>2018</year>) <volume>48</volume>(<issue>6</issue>):<page-range>1518&#x2013;28</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1002/jmri.26010</pub-id>
</citation>
</ref>
<ref id="B39">
<label>39</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lo</surname> <given-names>CM</given-names>
</name>
<name>
<surname>Chen</surname> <given-names>YC</given-names>
</name>
<name>
<surname>Weng</surname> <given-names>RC</given-names>
</name>
<name>
<surname>Hsieh</surname> <given-names>KLC</given-names>
</name>
</person-group>. <article-title>Intelligent Glioma Grading Based on Deep Transfer Learning of MRI Radiomic Features</article-title>. <source>Appl Sci-Basel</source> (<year>2019</year>) <volume>9</volume>(<issue>22</issue>). doi:&#xa0;ARTN 4926 10.3390/app9224926
</citation>
</ref>
<ref id="B40">
<label>40</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kumar</surname> <given-names>R</given-names>
</name>
<name>
<surname>Gupta</surname> <given-names>A</given-names>
</name>
<name>
<surname>Arora</surname> <given-names>HS</given-names>
</name>
<name>
<surname>Pandian</surname> <given-names>GN</given-names>
</name>
<name>
<surname>Raman</surname> <given-names>B</given-names>
</name>
</person-group>. <article-title>CGHF: A Computational Decision Support System for Glioma Classification Using Hybrid Radiomics- and Stationary Wavelet-Based Features</article-title>. <publisher-name>IEEE Access</publisher-name>(<year>2020</year>) <volume>8</volume>:<page-range>79440&#x2013;58</page-range>. doi:&#xa0;ARTN 4926 10.1109/Access.2020.2989193
</citation>
</ref>
<ref id="B41">
<label>41</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Devos</surname> <given-names>A</given-names>
</name>
<name>
<surname>Simonetti</surname> <given-names>AW</given-names>
</name>
<name>
<surname>van der Graaf</surname> <given-names>M</given-names>
</name>
<name>
<surname>Lukas</surname> <given-names>L</given-names>
</name>
<name>
<surname>Suykens</surname> <given-names>JAK</given-names>
</name>
<name>
<surname>Vanhamme</surname> <given-names>L</given-names>
</name>
<etal/>
</person-group>. <article-title>The Use of Multivariate MR Imaging Intensities Versus Metabolic Data From MR Spectroscopic Imaging for Brain Tumour Classification</article-title>. <source>J Magn Reson</source> (<year>2005</year>) <volume>173</volume>(<issue>2</issue>):<page-range>218&#x2013;28</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.jmr.2004.12.007</pub-id>
</citation>
</ref>
<ref id="B42">
<label>42</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ge</surname> <given-names>C</given-names>
</name>
<name>
<surname>Gu</surname> <given-names>IY</given-names>
</name>
<name>
<surname>Jakola</surname> <given-names>AS</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>J</given-names>
</name>
</person-group>. <article-title>Deep Learning and Multi-Sensor Fusion for Glioma Classification Using Multistream 2d Convolutional Networks</article-title>. <source>Annu Int Conf IEEE Eng Med Biol Soc</source> (<year>2018</year>) <volume>2018</volume>:<page-range>5894&#x2013;7</page-range>. doi: <pub-id pub-id-type="doi">10.1109/EMBC.2018.8513556</pub-id>
</citation>
</ref>
<ref id="B43">
<label>43</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Choy</surname> <given-names>G</given-names>
</name>
<name>
<surname>Khalilzadeh</surname> <given-names>O</given-names>
</name>
<name>
<surname>Michalski</surname> <given-names>M</given-names>
</name>
<name>
<surname>Do</surname> <given-names>S</given-names>
</name>
<name>
<surname>Samir</surname> <given-names>AE</given-names>
</name>
<name>
<surname>Pianykh</surname> <given-names>OS</given-names>
</name>
<etal/>
</person-group>. <article-title>Current Applications and Future Impact of Machine Learning in Radiology</article-title>. <source>Radiology</source> (<year>2018</year>) <volume>288</volume>(<issue>2</issue>):<page-range>318&#x2013;28</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1148/radiol.2018171820</pub-id>
</citation>
</ref>
<ref id="B44">
<label>44</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Louis</surname> <given-names>DN</given-names>
</name>
<name>
<surname>Perry</surname> <given-names>A</given-names>
</name>
<name>
<surname>Reifenberger</surname> <given-names>G</given-names>
</name>
<name>
<surname>von Deimling</surname> <given-names>A</given-names>
</name>
<name>
<surname>Figarella-Branger</surname> <given-names>D</given-names>
</name>
<name>
<surname>Cavenee</surname> <given-names>WK</given-names>
</name>
<etal/>
</person-group>. <article-title>The 2016 World Health Organization Classification of Tumors of the Central Nervous System: A Summary</article-title>. <source>Acta Neuropathol</source> (<year>2016</year>) <volume>131</volume>(<issue>6</issue>):<page-range>803&#x2013;20</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s00401-016-1545-1</pub-id>
</citation>
</ref>
<ref id="B45">
<label>45</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Brat</surname> <given-names>DJ</given-names>
</name>
<name>
<surname>Aldape</surname> <given-names>K</given-names>
</name>
<name>
<surname>Colman</surname> <given-names>H</given-names>
</name>
<name>
<surname>Figrarella-Branger</surname> <given-names>D</given-names>
</name>
<name>
<surname>Fuller</surname> <given-names>GN</given-names>
</name>
<name>
<surname>Giannini</surname> <given-names>C</given-names>
</name>
<etal/>
</person-group>. <article-title>cIMPACT-NOW Update 5: Recommended Grading Criteria and Terminologies for IDH-Mutant Astrocytomas</article-title>. <source>Acta Neuropathol</source> (<year>2020</year>) <volume>139</volume>(<issue>3</issue>):<page-range>603&#x2013;8</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s00401-020-02127-9</pub-id>
</citation>
</ref>
<ref id="B46">
<label>46</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Weller</surname> <given-names>M</given-names>
</name>
<name>
<surname>van den Bent</surname> <given-names>M</given-names>
</name>
<name>
<surname>Preusser</surname> <given-names>M</given-names>
</name>
<name>
<surname>Le Rhun</surname> <given-names>E</given-names>
</name>
<name>
<surname>Tonn</surname> <given-names>JC</given-names>
</name>
<name>
<surname>Minniti</surname> <given-names>G</given-names>
</name>
<etal/>
</person-group>. <article-title>EANO Guidelines on the Diagnosis and Treatment of Diffuse Gliomas of Adulthood</article-title>. <source>Nat Rev Clin Oncol</source> (<year>2021</year>) <volume>18</volume>(<issue>3</issue>):<page-range>170&#x2013;86</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/s41571-020-00447-z</pub-id>
</citation>
</ref>
<ref id="B47">
<label>47</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Maggio</surname> <given-names>I</given-names>
</name>
<name>
<surname>Franceschi</surname> <given-names>E</given-names>
</name>
<name>
<surname>Gatto</surname> <given-names>L</given-names>
</name>
<name>
<surname>Tosoni</surname> <given-names>A</given-names>
</name>
<name>
<surname>Di Nunno</surname> <given-names>V</given-names>
</name>
<name>
<surname>Tonon</surname> <given-names>C</given-names>
</name>
<etal/>
</person-group>. <article-title>Radiomics, Mirnomics, and Radiomirrnomics in Glioblastoma: Defining Tumor Biology From Shadow to Light</article-title>. <source>Expert Rev Anticancer Ther</source> (<year>2021</year>) <volume>21</volume>:<page-range>1265&#x2013;72</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1080/14737140.2021.1971518</pub-id>
</citation>
</ref>
<ref id="B48">
<label>48</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Whiting</surname> <given-names>PF</given-names>
</name>
<name>
<surname>Rutjes</surname> <given-names>AW</given-names>
</name>
<name>
<surname>Westwood</surname> <given-names>ME</given-names>
</name>
<name>
<surname>Mallett</surname> <given-names>S</given-names>
</name>
<name>
<surname>Deeks</surname> <given-names>JJ</given-names>
</name>
<name>
<surname>Reitsma</surname> <given-names>JB</given-names>
</name>
<etal/>
</person-group>. <article-title>QUADAS-2: A Revised Tool for the Quality Assessment of Diagnostic Accuracy Studies</article-title>. <source>Ann Intern Med</source> (<year>2011</year>) <volume>155</volume>(<issue>8</issue>):<page-range>529&#x2013;36</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.7326/0003-4819-155-8-201110180-00009</pub-id>
</citation>
</ref>
<ref id="B49">
<label>49</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mongan</surname> <given-names>J</given-names>
</name>
<name>
<surname>Moy</surname> <given-names>L</given-names>
</name>
<name>
<surname>Kahn</surname> <given-names>CE</given-names>
<suffix>Jr.</suffix>
</name>
</person-group> <article-title>Checklist for Artificial Intelligence in Medical Imaging (CLAIM): A Guide for Authors and Reviewers</article-title>. <source>Radiol Artif Intell</source> (<year>2020</year>) <volume>2</volume>(<issue>2</issue>):<elocation-id>e200029</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1148/ryai.2020200029</pub-id>
</citation>
</ref>
<ref id="B50">
<label>50</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Saito</surname> <given-names>T</given-names>
</name>
<name>
<surname>Rehmsmeier</surname> <given-names>M</given-names>
</name>
</person-group>. <article-title>The Precision-Recall Plot is More Informative Than the ROC Plot When Evaluating Binary Classifiers on Imbalanced Datasets</article-title>. <source>PloS One</source> (<year>2015</year>) <volume>10</volume>(<issue>3</issue>):<elocation-id>e0118432</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1371/journal.pone.0118432</pub-id>
</citation>
</ref>
<ref id="B51">
<label>51</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Knowler</surname> <given-names>WC</given-names>
</name>
<name>
<surname>Pettitt</surname> <given-names>DJ</given-names>
</name>
<name>
<surname>Savage</surname> <given-names>PJ</given-names>
</name>
<name>
<surname>Bennett</surname> <given-names>PH</given-names>
</name>
</person-group>. <article-title>Diabetes Incidence in Pima-Indians - Contributions of Obesity and Parental Diabetes</article-title>. <source>Am J Epidemiol</source> (<year>1981</year>) <volume>113</volume>(<issue>2</issue>):<page-range>144&#x2013;56</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1093/oxfordjournals.aje.a113079</pub-id>
</citation>
</ref>
<ref id="B52">
<label>52</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>DC</given-names>
</name>
<name>
<surname>Liu</surname> <given-names>CW</given-names>
</name>
<name>
<surname>Hu</surname> <given-names>SC</given-names>
</name>
</person-group>. <article-title>A Learning Method for the Class Imbalance Problem With Medical Data Sets</article-title>. <source>Comput Biol Med</source> (<year>2010</year>) <volume>5)</volume>:<page-range>509&#x2013;18</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compbiomed.2010.03.005</pub-id>
</citation>
</ref>
<ref id="B53">
<label>53</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ling</surname> <given-names>CX</given-names>
</name>
<name>
<surname>Huang</surname> <given-names>J</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>H</given-names>
</name>
</person-group>. <article-title>AUC: A Better Measure Than Accuracy in Comparing Learning Algorithms</article-title>. <source>Lect Notes Artif Int</source> (<year>2003</year>) <volume>2671</volume>:<page-range>329&#x2013;41</page-range>. doi: <pub-id pub-id-type="doi">10.1007/3-540-44886-1_25</pub-id>
</citation>
</ref>
<ref id="B54">
<label>54</label>
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cronin</surname> <given-names>P</given-names>
</name>
<name>
<surname>Kelly</surname> <given-names>AM</given-names>
</name>
<name>
<surname>Altaee</surname> <given-names>D</given-names>
</name>
<name>
<surname>Foerster</surname> <given-names>B</given-names>
</name>
<name>
<surname>Petrou</surname> <given-names>M</given-names>
</name>
<name>
<surname>Dwamena</surname> <given-names>BA</given-names>
</name>
</person-group>. <article-title>How to Perform a Systematic Review and Meta-Analysis of Diagnostic Imaging Studies</article-title>. <source>Acad Radiol</source> (<year>2018</year>) <volume>25</volume>(<issue>5</issue>):<page-range>573&#x2013;93</page-range>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.acra.2017.12.007</pub-id>
</citation>
</ref>
<ref id="B55">
<label>55</label>
<citation citation-type="web">
<person-group person-group-type="author">
<name>
<surname>Macaskill</surname> <given-names>P</given-names>
</name>
<name>
<surname>Gatsonis</surname> <given-names>C</given-names>
</name>
<name>
<surname>Deeks</surname> <given-names>JJ</given-names>
</name>
<name>
<surname>Harbord</surname> <given-names>RM</given-names>
</name>
<name>
<surname>Takwoingi</surname> <given-names>Y</given-names>
</name>
<collab>The Cochrane Collaboration</collab>
</person-group>. <article-title>Chapter 10: Analysing and Presenting Results</article-title>. In: <person-group person-group-type="editor">
<name>
<surname>Deeks</surname> <given-names>JJ</given-names>
</name>
<name>
<surname>Bossuyt</surname> <given-names>PM</given-names>
</name>
<name>
<surname>Gatsonis</surname> <given-names>C</given-names>
</name>
</person-group>, editors. <source>Cochrane Handbook for Systematic Reviews of Diagnostic Test Accuracy Version 1.0</source> <publisher-loc>Birmingham, UK</publisher-loc>: <publisher-name>The Cochrane Collaboration</publisher-name> (<year>2010</year>). Available at: <uri xlink:href="http://srdta.cochrane.org/">http://srdta.cochrane.org/</uri>.</citation>
</ref>
</ref-list>
</back>
</article>