<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD JATS (Z39.96) Journal Publishing DTD v1.3 20210610//EN" "JATS-journalpublishing1-3-mathml3.dtd">
<article article-type="research-article" dtd-version="1.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:ali="http://www.niso.org/schemas/ali/1.0/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Bioinform.</journal-id>
<journal-title-group>
<journal-title>Frontiers in Bioinformatics</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Bioinform.</abbrev-journal-title>
</journal-title-group>
<issn pub-type="epub">2673-7647</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1755664</article-id>
<article-id pub-id-type="doi">10.3389/fbinf.2026.1755664</article-id>
<article-version article-version-type="Version of Record" vocab="NISO-RP-8-2008"/>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Original Research</subject>
</subj-group>
</article-categories>
<title-group>
<article-title>Benchmarking multiple gene ontology enrichment tools reveals high biological significance, ranking, and stringency heterogeneity among datasets</article-title>
<alt-title alt-title-type="left-running-head">Oliveira et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/fbinf.2026.1755664">10.3389/fbinf.2026.1755664</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>de Oliveira</surname>
<given-names>F&#xe1;bio Henrique Schuster</given-names>
</name>
<xref ref-type="aff" rid="aff1"/>
<uri xlink:href="https://loop.frontiersin.org/people/2908023"/>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Conceptualization" vocab-term-identifier="https://credit.niso.org/contributor-roles/conceptualization/">Conceptualization</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Data curation" vocab-term-identifier="https://credit.niso.org/contributor-roles/data-curation/">Data curation</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Formal analysis" vocab-term-identifier="https://credit.niso.org/contributor-roles/formal-analysis/">Formal Analysis</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Investigation" vocab-term-identifier="https://credit.niso.org/contributor-roles/investigation/">Investigation</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Methodology" vocab-term-identifier="https://credit.niso.org/contributor-roles/methodology/">Methodology</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Validation" vocab-term-identifier="https://credit.niso.org/contributor-roles/validation/">Validation</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Visualization" vocab-term-identifier="https://credit.niso.org/contributor-roles/visualization/">Visualization</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; original draft" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-original-draft/">Writing - original draft</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &#x26; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/">Writing - review and editing</role>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Gomes</surname>
<given-names>Felipe Acker</given-names>
</name>
<xref ref-type="aff" rid="aff1"/>
<uri xlink:href="https://loop.frontiersin.org/people/3358071"/>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Formal analysis" vocab-term-identifier="https://credit.niso.org/contributor-roles/formal-analysis/">Formal Analysis</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Validation" vocab-term-identifier="https://credit.niso.org/contributor-roles/validation/">Validation</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &#x26; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/">Writing - review and editing</role>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Feltes</surname>
<given-names>Bruno C&#xe9;sar</given-names>
</name>
<xref ref-type="aff" rid="aff1"/>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/490457"/>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Conceptualization" vocab-term-identifier="https://credit.niso.org/contributor-roles/conceptualization/">Conceptualization</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Data curation" vocab-term-identifier="https://credit.niso.org/contributor-roles/data-curation/">Data curation</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Formal analysis" vocab-term-identifier="https://credit.niso.org/contributor-roles/formal-analysis/">Formal Analysis</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Funding acquisition" vocab-term-identifier="https://credit.niso.org/contributor-roles/funding-acquisition/">Funding acquisition</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Investigation" vocab-term-identifier="https://credit.niso.org/contributor-roles/investigation/">Investigation</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Methodology" vocab-term-identifier="https://credit.niso.org/contributor-roles/methodology/">Methodology</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Project administration" vocab-term-identifier="https://credit.niso.org/contributor-roles/project-administration/">Project administration</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Resources" vocab-term-identifier="https://credit.niso.org/contributor-roles/resources/">Resources</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Supervision" vocab-term-identifier="https://credit.niso.org/contributor-roles/supervision/">Supervision</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Visualization" vocab-term-identifier="https://credit.niso.org/contributor-roles/visualization/">Visualization</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; original draft" vocab-term-identifier="https://credit.niso.org/contributor-roles/writing-original-draft/">Writing - original draft</role>
<role vocab="credit" vocab-identifier="https://credit.niso.org/" vocab-term="Writing &#x2013; review &#x26; editing" vocab-term-identifier="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/">Writing - review and editing</role>
</contrib>
</contrib-group>
<aff id="aff1">
<institution>Laboratory of DNA Repair and Aging, Department of Biophysics, Institute of Biosciences, Federal University of Rio Grande do Sul</institution>, <city>Porto Alegre</city>, <country country="BR">Brazil</country>
</aff>
<author-notes>
<corresp id="c001">
<label>&#x2a;</label>Correspondence: Bruno C&#xe9;sar Feltes, <email xlink:href="mailto:bruno.feltes@ufrgs.br">bruno.feltes@ufrgs.br</email>
</corresp>
</author-notes>
<pub-date publication-format="electronic" date-type="pub" iso-8601-date="2026-01-29">
<day>29</day>
<month>01</month>
<year>2026</year>
</pub-date>
<pub-date publication-format="electronic" date-type="collection">
<year>2026</year>
</pub-date>
<volume>6</volume>
<elocation-id>1755664</elocation-id>
<history>
<date date-type="received">
<day>27</day>
<month>11</month>
<year>2025</year>
</date>
<date date-type="rev-recd">
<day>10</day>
<month>01</month>
<year>2026</year>
</date>
<date date-type="accepted">
<day>12</day>
<month>01</month>
<year>2026</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2026 de Oliveira, Gomes and Feltes.</copyright-statement>
<copyright-year>2026</copyright-year>
<copyright-holder>de Oliveira, Gomes and Feltes</copyright-holder>
<license>
<ali:license_ref start_date="2026-01-29">https://creativecommons.org/licenses/by/4.0/</ali:license_ref>
<license-p>This is an open-access article distributed under the terms of the <ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">Creative Commons Attribution License (CC BY)</ext-link>. The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</license-p>
</license>
</permissions>
<abstract>
<p>Functional enrichment analysis (FEA) provides biological meaning from lists of differentially expressed genes and proteins obtained through omics experiments. FEA tools can employ numerous statistical methods and rely on different pathway databases. In this sense, Overrepresentation Analysis (ORA) is one of the most popular methods to perform FEA. Gene Ontology (GO) is arguably the most widely used pathway knowledgebase in FEA. Hence, benchmarking the biological accuracy of ORA-based GO enrichment tools is crucial. Nevertheless, benchmark studies in FEA tend to focus excessively on performance-based metrics rather than on the biological information contained in enrichment results. To identify the differences between popular ORA-based GO enrichment tools and provide data that brings insights into the tools&#x2019; biological accuracy and, thus, better suits the application of FEA, we tested 12 popular GO enrichment tools (i.e., DAVID, PANTHER, WebGestalt, Enrichr, ShinyGO, limma, topGO, GOstats, clusterProfiler, g:Profiler, ClueGO, and BiNGO) with randomized datasets as negative controls, a target-oriented and a hallmark datasets as positive controls, and an experiment-derived dataset. Gene sets with 500, 200, 100, and 50 genes were built for each dataset to investigate the impact of input sizes. Using the control datasets, we calculated the FPR and accuracy of the tools based on the semantic similarity between the enriched terms and the target ontologies and assessed overlooked, insightful metrics that reflect the biological informativeness of the results, such as the specificity of enriched GO terms and the prioritization of target ontologies. Additionally, we clustered the FEA results based on term semantic similarity, enabling us to directly compare the biological profiles generated by each tool. Despite employing the same method and functional database, the tools&#x2019; results diverged significantly. Our findings reveal considerable variation among tools in terms of informativeness and interpretability of results. Some tools demonstrated strong capabilities in prioritizing target pathways, while others struggled, especially as input size increased. Additionally, we observed that the degree to which the enriched ontologies are related to the expected targets varies across tools, with some being more conservative than others. Together, these results provide powerful insights into the performance characteristics of the analyzed GO enrichment tools and yield new, relevant data for benchmarking FEA tools.</p>
</abstract>
<kwd-group>
<kwd>benchmark</kwd>
<kwd>bioinformatics</kwd>
<kwd>functional enrichment</kwd>
<kwd>gene</kwd>
<kwd>ontology</kwd>
<kwd>overrepresentation analysis</kwd>
</kwd-group>
<funding-group>
<award-group id="gs1">
<funding-source id="sp1">
<institution-wrap>
<institution>Funda&#xe7;&#xe3;o de Amparo &#xe0; Pesquisa do Estado do Rio Grande do Sul</institution>
<institution-id institution-id-type="doi" vocab="open-funder-registry" vocab-identifier="10.13039/open_funder_registry">10.13039/501100004263</institution-id>
</institution-wrap>
</funding-source>
<award-id rid="sp1">24/2551-0001277-0</award-id>
</award-group>
<award-group id="gs2">
<funding-source id="sp2">
<institution-wrap>
<institution>Coordena&#xe7;&#xe3;o de Aperfei&#xe7;oamento de Pessoal de N&#xed;vel Superior</institution>
<institution-id institution-id-type="doi" vocab="open-funder-registry" vocab-identifier="10.13039/open_funder_registry">10.13039/501100002322</institution-id>
</institution-wrap>
</funding-source>
<award-id rid="sp2">001</award-id>
</award-group>
<funding-statement>The author(s) declared that financial support was received for this work and/or its publication. This work was funded by the Coordena&#xe7;&#xe3;o de Aperfei&#xe7;oamento de Pessoal de N&#xed;vel Superior (CAPES, Brazil; Finance code 001) and the Funda&#xe7;&#xe3;o de Amparo &#xe0; Pesquisa do Estado do Rio Grande do Sul (FAPERGS) [24/2551-0001277-0].</funding-statement>
</funding-group>
<counts>
<fig-count count="6"/>
<table-count count="4"/>
<equation-count count="0"/>
<ref-count count="47"/>
<page-count count="00"/>
</counts>
<custom-meta-group>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Integrative Bioinformatics</meta-value>
</custom-meta>
</custom-meta-group>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<label>1</label>
<title>Introduction</title>
<p>Functional enrichment analysis (FEA) is a widely used method that provides additional biological meaning from lists of differentially expressed genes (DEGs) and proteins obtained primarily from high-throughput omics experiments by identifying enriched &#x201c;functional descriptions&#x201d; within omics data. FEA tools use the knowledge contained in functional databases, which associate functional categories with gene lists, such as the Gene Ontology (GO) knowledgebase (<xref ref-type="bibr" rid="B3">Ashburner et al., 2000</xref>; <xref ref-type="bibr" rid="B10">Consortium et al., 2023</xref>), the Kyoto Encyclopedia of Genes and Genomes (KEGG) (<xref ref-type="bibr" rid="B21">Kanehisa and Goto, 2000</xref>; <xref ref-type="bibr" rid="B22">Kanehisa et al., 2025</xref>), WikiPathways (<xref ref-type="bibr" rid="B1">Agrawal et al., 2024</xref>), and Reactome (<xref ref-type="bibr" rid="B32">Milacic et al., 2024</xref>). In this scenario, the GO is one of the most widely used resources in the scientific community for providing functional information on genes and gene products.</p>
<p>Currently, a variety of methods that rely on different databases and statistical approaches have been developed to conduct FEA. Most available methods can be classified into four main classes: Overrepresentation Analysis (ORA), Functional Class Scoring (FCS), Pathway-topology-based (PT), and Network-based (NB). Due to their importance, FEA has become embedded in nearly all omics analysis protocols. However, despite employing the same method, different enrichment tools produce different outputs. Due to such inherent heterogeneity, efforts are made to benchmark the performance of distinct enrichment approaches and tools (<xref ref-type="bibr" rid="B39">Tarca et al., 2013</xref>; <xref ref-type="bibr" rid="B4">Bayerlov&#xe1; et al., 2015</xref>; <xref ref-type="bibr" rid="B28">Lim et al., 2018</xref>; <xref ref-type="bibr" rid="B33">Nguyen et al., 2019</xref>; <xref ref-type="bibr" rid="B47">Zyla et al., 2019</xref>; <xref ref-type="bibr" rid="B17">Geistlinger et al., 2021</xref>; <xref ref-type="bibr" rid="B8">Buzzao et al., 2024</xref>). Nevertheless, such benchmark studies tend to focus on comparing the statistical methods (i.e., ORA, FCS, PT, and PT), instead of comparing the results profile generated by them (<xref ref-type="bibr" rid="B39">Tarca et al., 2013</xref>; <xref ref-type="bibr" rid="B28">Lim et al., 2018</xref>; <xref ref-type="bibr" rid="B17">Geistlinger et al., 2021</xref>; <xref ref-type="bibr" rid="B8">Buzzao et al., 2024</xref>). Moreover, by using only a few tools to represent a whole class (e.g., DAVID for ORA; GSEA for FCS), these studies neglect the differences among software based on the same method (<xref ref-type="bibr" rid="B11">Dong et al., 2016</xref>; <xref ref-type="bibr" rid="B33">Nguyen et al., 2019</xref>; <xref ref-type="bibr" rid="B47">Zyla et al., 2019</xref>). Another commonly overlooked limitation of benchmarks in the case of FEA is that the comparisons tend to rely solely on standard performance metrics, such as FDR, sensitivity, accuracy, and specificity, which fail to accurately illustrate the performance and behavior of FEA tools, as they disregard the biological information that the results provide.</p>
<p>In this study, we evaluated the behavior of 12 commonly used ORA-based tools (<xref ref-type="table" rid="T1">Table 1</xref>) that utilize the GO resource for FEA. To evaluate the performance of the selected tools, we used an approach focused on the biological meaning derived from the FEA. We conducted enrichment analysis using <italic>random</italic> datasets, a <italic>Hallmark</italic> dataset, a GO Biological Process (<italic>GOBP</italic>) gene set, and a microarray-derived dataset, all split into lists of varying sizes. The random group served as a negative control, while the hallmark dataset was employed as a positive control. Furthermore, the <italic>GOBP</italic> dataset was constructed with predetermined target ontologies to enable the calculation of relevant metrics, including accuracy and FPR, and to assess the tools&#x2019; ranking abilities of the target pathways. Finally, we constructed the <italic>Contextual</italic> lists using real high-throughput experiment data to evaluate the differences in the results of various tools in a realistic research scenario. We also used GO term annotation size and depth in the ontology as measures for biological specificity. Such an approach has been used in previous studies, but is not commonly employed in FEA benchmarks (<xref ref-type="bibr" rid="B25">Lewin and Grieve, 2006</xref>; <xref ref-type="bibr" rid="B30">Louie et al., 2010</xref>; <xref ref-type="bibr" rid="B40">Tomczak et al., 2018</xref>).</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Selected tools for the comparative analysis and relevant information. The GO version column corresponds to the GO version used at the time of the analyses according to each tool&#x2019;s documentation. The column &#x201c;Raw p-value&#x201d; indicates whether the tool also provides raw p-values alongside their corrected values.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Software</th>
<th align="center">GO version</th>
<th align="center">Custom annotation and GO files</th>
<th align="center">Raw p-value</th>
<th align="center">Platform</th>
<th align="center">Release</th>
<th align="center">Version/Last updated</th>
<th align="center">Reference</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">DAVID</td>
<td align="center">2025<sup>a</sup>
</td>
<td align="center">No/No</td>
<td align="center">Yes</td>
<td align="center">Web</td>
<td align="center">2003</td>
<td align="right">4/1/2024 (DAVID Knowledgebase v2024q1)</td>
<td align="right">
<xref ref-type="bibr" rid="B18">Huang et al., 2009</xref>; <xref ref-type="bibr" rid="B37">Sherman et al., 2022</xref>
</td>
</tr>
<tr>
<td align="left">PANTHER</td>
<td align="center">2025-02-06</td>
<td align="center">No/No</td>
<td align="center">Yes</td>
<td align="center">Web</td>
<td align="center">2000</td>
<td align="right">v18.0 - 17/09/2023</td>
<td align="right">
<xref ref-type="bibr" rid="B31">Mi et al. (2019)</xref>
</td>
</tr>
<tr>
<td align="left">Enrichr</td>
<td align="center">2025<sup>b</sup>
</td>
<td align="center">No/No</td>
<td align="center">Yes</td>
<td align="center">Web</td>
<td align="center">2013</td>
<td align="right">8/7/2023</td>
<td align="right">
<xref ref-type="bibr" rid="B9">Chen et al. (2013)</xref>
</td>
</tr>
<tr>
<td align="left">WebGestalt</td>
<td align="center">2024<sup>b</sup>
</td>
<td align="center">No/No</td>
<td align="center">Yes</td>
<td align="center">Web/R package</td>
<td align="center">2005</td>
<td align="right">2024</td>
<td align="right">
<xref ref-type="bibr" rid="B14">Elizarraras et al. (2024)</xref>
</td>
</tr>
<tr>
<td align="left">g:Profiler</td>
<td align="center">2024<sup>b</sup>
</td>
<td align="center">Yes/No</td>
<td align="center">No</td>
<td align="center">Web/R package</td>
<td align="center">2007</td>
<td align="right">e112_eg59_p19_25aa4782 - 2025</td>
<td align="right">
<xref ref-type="bibr" rid="B24">Kolberg et al. (2023)</xref>
</td>
</tr>
<tr>
<td align="left">ShinyGO</td>
<td align="center">2022<sup>b</sup>
</td>
<td align="center">No/No</td>
<td align="center">No</td>
<td align="center">Web</td>
<td align="center">2018</td>
<td align="right">v0.82 - 2/2025</td>
<td align="right">
<xref ref-type="bibr" rid="B16">Ge et al. (2020)</xref>
</td>
</tr>
<tr>
<td align="left">ClueGO</td>
<td align="center">2025-03-16</td>
<td align="center">Yes/Yes</td>
<td align="center">Yes</td>
<td align="center">Cytoscape</td>
<td align="center">2009</td>
<td align="right">v2.5.10 - 2023</td>
<td align="right">
<xref ref-type="bibr" rid="B6">Bindea et al. (2009)</xref>
</td>
</tr>
<tr>
<td align="left">BiNGO</td>
<td align="center">2013<sup>b</sup>
</td>
<td align="center">Yes/Yes</td>
<td align="center">Yes</td>
<td align="center">Cytoscape</td>
<td align="center">2005</td>
<td align="right">v3.0.5 - 2021</td>
<td align="right">
<xref ref-type="bibr" rid="B51">Maere et al. (2005)</xref>
</td>
</tr>
<tr>
<td align="left">topGO</td>
<td align="center">2024-09-20</td>
<td align="center">Yes/No</td>
<td align="center">Yes</td>
<td align="center">R package</td>
<td align="center">2006</td>
<td align="right">v2.58.0 - 2024</td>
<td align="right">
<xref ref-type="bibr" rid="B2">Alexa A (2024)</xref>
</td>
</tr>
<tr>
<td align="left">clusterProfiler</td>
<td align="center">2024-09-20</td>
<td align="center">Yes/No</td>
<td align="center">Yes</td>
<td align="center">R package</td>
<td align="center">2012</td>
<td align="right">v4.14.6 - 2025</td>
<td align="right">
<xref ref-type="bibr" rid="B45">Xu et al. (2024)</xref>
</td>
</tr>
<tr>
<td align="left">GOstats</td>
<td align="center">2024-09-20</td>
<td align="center">No/No</td>
<td align="center">Yes</td>
<td align="center">R package</td>
<td align="center">2006</td>
<td align="right">v2.72.0 - 2025</td>
<td align="right">
<xref ref-type="bibr" rid="B15">Falcon and Gentleman (2007)</xref>
</td>
</tr>
<tr>
<td align="left">goana (limma)</td>
<td align="center">2024-09-20</td>
<td align="center">No/No</td>
<td align="center">Yes</td>
<td align="center">R package</td>
<td align="center">2015</td>
<td align="right">Limma v3.62.2 - 01/2025</td>
<td align="right">
<xref ref-type="bibr" rid="B35">Ritchie et al. (2015)</xref>
</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<fn>
<p>Table legends: <sup>a</sup> updated daily; <sup>b</sup> only the update year is mentioned.</p>
</fn>
</table-wrap-foot>
</table-wrap>
<p>In addition to providing a comprehensive overview of how leading ORA-based tools perform in terms of output consistency and biological relevance, our work also proposes a novel benchmarking strategy to guide the evaluation of future tools in the field. The goal of this work is not to analyze the performance of the ORA tools, but to examine the coherence and profile of the biological information they provide across different inputs, list sizes, and their precision in ranking and identifying expected results.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<label>2</label>
<title>Materials and methods</title>
<sec id="s2-1">
<label>2.1</label>
<title>Dataset generation</title>
<p>For each dataset, we created lists with 500, 200, 100, and 50 genes. Furthermore, Entrez ID gene lists were created for all lists by converting their gene symbols using the biomaRt library in R (<xref ref-type="bibr" rid="B12">Durinck et al., 2005</xref>; <xref ref-type="bibr" rid="B13">Durinck et al., 2009</xref>). Both lists, with gene symbols and Entrez IDs, were used because some tools either accepted or displayed imprecise results for one of the identifiers. Our datasets can be acquired from [<ext-link ext-link-type="uri" xlink:href="https://github.com/LARA-Lab-Aging">https://github.com/LARA-Lab-Aging</ext-link>].</p>
<p>The random lists were generated by selecting 500, 200, 100, and 50 protein-coding <italic>Homo sapiens</italic> genes at random. This process was repeated 5 times to produce 5 different lists for each size.</p>
<p>The <italic>Hallmark</italic> lists were built from the gene sets available in the MSigDB human hallmarks collection (<xref ref-type="bibr" rid="B38">Subramanian et al., 2005</xref>; <xref ref-type="bibr" rid="B26">Liberzon et al., 2011</xref>; <xref ref-type="bibr" rid="B27">2015</xref>). To create the 500 genes <italic>Hallmark</italic> list, we combined the HALLMARK HYPOXIA, HALLMARK DNA REPAIR, and HALLMARK UV RESPONSE UP gene sets. The 200-gene list contained only genes in the hypoxia set. The 100 and 50-gene lists were built by downsampling the 200-hypoxia dataset.</p>
<p>The <italic>GOBP</italic> lists were generated by combining gene sets from different GO Biological Process datasets in the MSigDB Ontologies collection. To build the lists, we aimed to select ontologies with minimal overlap. The ontology gene sets that compose each gene list are gathered in <xref ref-type="sec" rid="s11">Supplementary Table S1</xref>. Additionally, all lists had their duplicates removed.</p>
<p>The <italic>Contextual</italic> dataset was obtained from the lung cancer GSE18842 dataset available in the Gene Expression Omnibus (GEO) [<ext-link ext-link-type="uri" xlink:href="https://www.ncbi.nlm.nih.gov/geo/">https://www.ncbi.nlm.nih.gov/geo/</ext-link>], which comprises 46 cancer and 45 control samples (<xref ref-type="bibr" rid="B48">Sanchez-Palencia et al., 2011</xref>). Firstly, the dataset was imported into Gene Expression Analysis Platform (GEAP) software, which employs underlying R-based tools through a graphical user interface to perform microarray data analyses (<xref ref-type="bibr" rid="B34">Nunes et al., 2022</xref>). Quality analysis was conducted to filter out low-quality data within GEAP, which utilizes the <italic>arrayQualityMetrics</italic> package internally. Samples that failed the quality metrics of at least two of the three metrics analyzed by <italic>arrayQualityMetrics</italic> would be discarded before the Differential Gene Expression (DGE) analysis; however, no samples had to be discarded. Next, DGE analysis was conducted through the &#x201c;Comparison Between Two Groups&#x201d; tab in GEAP, which employs the limma package. The parameters used were the default eBayes method and the False Discovery Rate (FDR) correction method (<xref ref-type="bibr" rid="B5">Benjamini and Hochberg, 1995</xref>). Finally, the results were filtered for logFC &#x3e;1 and p-value &#x3c;0.05, yielding a table with 3,222 DEG (overexpressed &#x3d; 1,386; underexpressed &#x3d; 1836). The top 500 DEGs were selected to build the 500 contextual gene list. The same downsampling process as described before was used to create the smaller lists.</p>
</sec>
<sec id="s2-2">
<label>2.2</label>
<title>Tool selection and FEA</title>
<p>Tools were selected based on the following criteria: (i) tools widely used in the scientific community (&#x2265;500 citations on Google Scholar), (ii) tools that have been updated in the past 5 years or allow the usage of up-to-date GO annotations and ontology files, and (iii) tools that are functioning as described in their available documentation. For the FEA, we attempted to maintain the parameters as similar as possible and used the default options when no equivalents were available. All analyses were conducted using human GO annotations. The whole human genome was used as background for the <italic>Random</italic>, <italic>Hallmark</italic>, and <italic>GOBP</italic> sets, whereas the Affymetrix table from GSE18842 served as background for the <italic>Contextual</italic> dataset. In general, either the FDR or the Benjamini-Hochberg (BH) correction method was used, and all ontologies with fewer than 2 genes were removed from the results. Enrichment results for the <italic>Hallmark</italic>, <italic>Contextual</italic>, and <italic>GOBP</italic> datasets are compiled in <xref ref-type="sec" rid="s11">Supplementary Table S2</xref>. All raw outputs are available at [<ext-link ext-link-type="uri" xlink:href="https://github.com/LARA-Lab-Aging">https://github.com/LARA-Lab-Aging</ext-link>]. Additionally, gene-mapping success rates are provided in <xref ref-type="sec" rid="s11">Supplementary Table S5</xref>.</p>
<sec id="s2-2-1">
<label>2.2.1</label>
<title>BiNGO</title>
<p>Analyses that employed BiNGO were used to assess overrepresentation and employed the hypergeometric test to compute p-values (<xref ref-type="bibr" rid="B51">Maere et al., 2005</xref>). Additionally, the significance level was set to 1, all categories were selected, and the ontology file used was GO_Biological_Process in BiNGO.</p>
</sec>
<sec id="s2-2-2">
<label>2.2.2</label>
<title>ClueGO</title>
<p>ClueGO enrichment results were obtained through the One-sided hypergeometric test (enrichment) with the GO Biological Process annotation set, the evidence option set to &#x201c;All&#x201d;, minimum number of genes per ontology set to two, and the BH option for p-value correction (<xref ref-type="bibr" rid="B6">Bindea et al., 2009</xref>). Other options available were all unselected.</p>
</sec>
<sec id="s2-2-3">
<label>2.2.3</label>
<title>DAVID</title>
<p>Results from DAVID were obtained by querying the 16 gene lists (Entrez IDs) in the functional annotation tab and selecting the GOTERM_BP_DIRECT chart (<xref ref-type="bibr" rid="B18">Huang et al., 2009</xref>; <xref ref-type="bibr" rid="B37">Sherman et al., 2022</xref>). The parameters used were EASE &#x2264;1, Count &#x2265;2, and the maximum number of records was set at 10,000. The correction method employed was FDR.</p>
</sec>
<sec id="s2-2-4">
<label>2.2.4</label>
<title>Enrichr</title>
<p>The results in Enrichr were obtained by querying the 16 Gene Symbol gene lists and selecting the GO Biological Process 2025 in the Ontologies tab (<xref ref-type="bibr" rid="B9">Chen et al., 2013</xref>). The correction method was BH.</p>
</sec>
<sec id="s2-2-5">
<label>2.2.5</label>
<title>GOstats</title>
<p>Analyses in the GOstatsR package were conducted through the hyperGTest function (annotation &#x3d; &#x201c;org.Hs.eg.db&#x201d;, ontology &#x3d; &#x201c;BP&#x201d;, pvalueCutoff &#x3d; 1, testDirection &#x3d; &#x201c;over&#x201d;) that is based on the hypergeometric distribution statistical test (<xref ref-type="bibr" rid="B15">Falcon and Gentleman, 2007</xref>).</p>
</sec>
<sec id="s2-2-6">
<label>2.2.6</label>
<title>PANTHER</title>
<p>Analyses in PANTHER were conducted with the Statistical Overrepresentation test and the Biological Process complete annotation set (<xref ref-type="bibr" rid="B31">Mi et al., 2019</xref>). The statistical test used was Fisher&#x2019;s exact test, and the correction method employed was FDR. PANTHER could not map Entrez IDs properly. Thus, lists were queried with the Gene Symbols.</p>
</sec>
<sec id="s2-2-7">
<label>2.2.7</label>
<title>ShinyGO</title>
<p>ShinyGO analyses results were generated in the Enrichment tool with the Pathway database option set to GO Biological Process, FDR cutoff to 1.0, and pathway minimal size to 2 (<xref ref-type="bibr" rid="B16">Ge et al., 2020</xref>). The redundancy removal option was unselected. Only the top 1,000 enriched ontologies were selected, which is the maximum number of results ShinyGO outputs.</p>
</sec>
<sec id="s2-2-8">
<label>2.2.8</label>
<title>WebGestalt</title>
<p>Analyses in WebGestalt were conducted by using the ORA method option with the Biological Process GO Functional Database (<xref ref-type="bibr" rid="B14">Elizarraras et al., 2024</xref>). No redundancy removal method was selected, and the FDR was used as the p-value adjustment method. The top 10,000 ontologies were selected as the results.</p>
</sec>
<sec id="s2-2-9">
<label>2.2.9</label>
<title>clusterProfiler</title>
<p>Enrichment results in clusterProfiler were obtained through the enrichGO function (OrgDB &#x3d; &#x201c;org.Hs.eg.db&#x201d;, ont &#x3d; &#x201c;BP&#x201d;, pAdjustMethod &#x3d; &#x201c;fdr&#x201d;, pvalueCutoff &#x3d; 1, minGSSize &#x3d; 2), which uses the hypergeometric distribution to conduct statistical analysis (<xref ref-type="bibr" rid="B45">Xu et al., 2024</xref>).</p>
</sec>
<sec id="s2-2-10">
<label>2.2.10</label>
<title>g:Profiler</title>
<p>The analyses with g:Profiler were conducted in the R package gprofiler2 through the gost function (organism &#x3d; &#x201c;hsapiens&#x201d;, significant &#x3d; F, user_threshold &#x3d; 1, correction_method &#x3d; &#x201c;fdr&#x201d;, sources &#x3d; &#x201c;GO:BP&#x201d;) (<xref ref-type="bibr" rid="B24">Kolberg et al., 2023</xref>).</p>
</sec>
<sec id="s2-2-11">
<label>2.2.11</label>
<title>goana (limma)</title>
<p>The goana function from the limma R package was used to conduct the analyses, with the FDR parameter set to 0.05 and the species set to &#x201c;Hs&#x201d; (<xref ref-type="bibr" rid="B35">Ritchie et al., 2015</xref>). The custom background option was used with the Contextual lists.</p>
</sec>
<sec id="s2-2-12">
<label>2.2.12</label>
<title>topGO</title>
<p>Analyses with topGO were conducted with the tool&#x2019;s default weight01 algorithm, Fisher&#x2019;s exact statistical test, org. Hs.eg.db annotation data, and the BP ontology set (<xref ref-type="bibr" rid="B2">Alexa, 2024</xref>). Subsequently, p-value correction with FDR through R&#x2019;s p. adjust function before removal of ontologies annotated to less than 2 genes in the input list to ensure adequate p-value correction.</p>
</sec>
</sec>
<sec id="s2-3">
<label>2.3</label>
<title>GO term specificity assessment</title>
<p>GO term biological specificity was assessed based on the term&#x2019;s annotation size and its depth in the ontology structure. Both properties have been widely employed to root the evaluation of term specificity (<xref ref-type="bibr" rid="B25">Lewin and Grieve, 2006</xref>; <xref ref-type="bibr" rid="B30">Louie et al., 2010</xref>; <xref ref-type="bibr" rid="B40">Tomczak et al., 2018</xref>). Additionally, the link between them and biological specificity is highly intuitive: a term tends to be more general as more genes are associated with it, and, because the ontology is hierarchical, terms deeper in the hierarchy tend to be more biologically precise. Metrics were retrieved using the GOATOOLS Python library (<xref ref-type="bibr" rid="B23">Klopfenstein et al., 2018</xref>).</p>
</sec>
<sec id="s2-4">
<label>2.4</label>
<title>Semantic similarity analysis</title>
<p>Semantic similarity (SS) analyses were conducted using GOATOOLS, which implements multiple methods to determine semantic similarity between two GO terms. We selected Wang&#x2019;s method (<xref ref-type="bibr" rid="B42">Wang et al., 2007</xref>) to conduct the SS analyses, as it relies solely on the GO Directed Acyclic Graph structure to define SS and attempts to translate the similarity of two ontologies into biological meaning. The GO&#x2019;s relationships &#x2018;is_a&#x2019; and &#x2018;part_of&#x2019; with edge scores of 0.8 and 0.6, respectively, were used to determine Wang&#x2019;s SS score.</p>
</sec>
<sec id="s2-5">
<label>2.5</label>
<title>Metrics calculation</title>
<p>To calculate the defined metrics, we used the enrichment results obtained with the <italic>GOBP</italic> dataset. True positives (TP) were defined as statistically significant ontologies (adjusted p-value &#x3c;0.05) that had a Wang&#x2019;s semantic similarity (SS) score of at least 0.7 with at least one of the target ontologies for that input list. Conversely, false negatives (FN) were ontologies with a Wang&#x2019;s SS score of at least 0.7, but that were not statistically significant. Similarly, false positives (FP) and true negatives (TN) were ontologies with low maximum semantic similarity (Wang&#x2019;s SS score &#x3c;0.3) compared to the original target pathways, which were statistically significant or not, respectively. These Wang&#x2019;s SS score thresholds were selected to center the analysis on the ontologies that are very similar (Wang&#x2019;s SS score &#x3e;0.7) to one of the targets and are, therefore, expected in the results; and on ontologies that are very dissimilar (Wang&#x2019;s SS score), thus are unexpected.</p>
</sec>
<sec id="s2-6">
<label>2.6</label>
<title>Enriched GO terms network construction and clustering</title>
<p>To group the enriched ontologies into functionally similar clusters, we built interaction networks for all FEA results, in which the edge score between two GO terms was their respective Wang&#x2019;s SS scores. The networks were then pruned with an edge-weight cutoff of 0.5 to remove edges connecting dissimilar ontologies. The resulting ontology similarity networks were clustered using Markov Clustering (MCL) (<xref ref-type="bibr" rid="B41">Van Dongen, 2008</xref>), one of the most robust algorithms for clustering biological data (<xref ref-type="bibr" rid="B7">Broh&#xe9;e and van Helden, 2006</xref>; <xref ref-type="bibr" rid="B36">Satuluri et al., 2010</xref>; <xref ref-type="bibr" rid="B29">Lim et al., 2019</xref>). Of the four Markov Clustering inflation values tested &#x2013; 1.5, 2.0, 3.5, and 5.0 &#x2013;, the 5.0 value yielded GO clusters with the highest biological accuracy. Finally, we classified the 15 largest clusters that contained 3 or more ontologies for all the clustered FEA results (<xref ref-type="sec" rid="s11">Supplementary Tables S3, S4</xref>).</p>
<p>Our pipeline is summarized in <xref ref-type="fig" rid="F1">Figure 1</xref>.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Procedure developed to evaluate the 12 ORA-based GO enrichment tools. Four different datasets were generated, each containing inputs with 500, 200, 100, and 50 genes. FEA results were then analyzed comparatively, and tools&#x2019; performance was assessed based on (i) the number of enriched ontologies, (ii) the biological informativeness of their results, (iii) target prioritization, (iv) stringency levels, and (v) biological profile.</p>
</caption>
<graphic xlink:href="fbinf-06-1755664-g001.tif">
<alt-text content-type="machine-generated">Diagram illustrating the workflow for functional enrichment analysis (FEA). It starts with the Molecular Signatures Database (MSigDB) and Gene Expression Omnibus (GEO) data, progressing through GOBP, Hallmark, and Contextual pathways, using various tools like BiNGO, DAVID, and PANTHER. Results undergo numerous evaluations: size, biological informativeness, target prioritization, stringency, and functional clustering, culminating in 192 FEA results.</alt-text>
</graphic>
</fig>
</sec>
</sec>
<sec sec-type="results" id="s3">
<label>3</label>
<title>Results</title>
<sec id="s3-1">
<label>3.1</label>
<title>Some tools identify statistically significant ontologies in random datasets</title>
<p>To evaluate how tools handle data with seemingly no biological context, we conducted FEA using multiple randomized lists of human protein-coding genes and analyzed the distributions of the number of enriched GO terms across input sizes. In general, when using the nominal p-value to determine significance, tools do yield enrichment results, and the number of enriched ontologies varies considerably across tools (<xref ref-type="sec" rid="s11">Supplementary Figure S1</xref>).</p>
<p>Nevertheless, after adjusting for multiple comparisons, most results for all tested tools are not statistically significant (p-value &#x3c;0.05) (<xref ref-type="fig" rid="F2">Figure 2</xref>), reinforcing the importance of analyzing enrichment results using corrected p-values, as this dramatically reduces the number of false positives, especially when employing ORA (<xref ref-type="bibr" rid="B19">Hung et al., 2012</xref>; <xref ref-type="bibr" rid="B43">Wijesooriya et al., 2022</xref>). However, there are tools able to retrieve enriched GO terms despite the nature of the dataset and the statistical correction of p-values (<xref ref-type="fig" rid="F2">Figure 2</xref>). Remarkably, ClueGO and goana consistently yielded the largest number of enriched categories for this dataset (<xref ref-type="fig" rid="F2">Figure 2</xref>). In particular, ClueGO displayed a rather interesting behavior: there are substantially fewer enriched GO terms for the 500-lists in comparison to the other input sizes, despite it retrieving larger results when considering nominal p-values (<xref ref-type="sec" rid="s11">Supplementary Figure S1</xref>). Such behavior could reflect how the p-value correction method is implemented and how detected but not enriched GO terms are handled. Ultimately, this suggests that both ClueGO and goana are susceptible to retrieving enriched GO terms that are either irrelevant or insufficiently descriptive.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Distribution of the number of enriched GO terms for the input lists in the <italic>Random</italic> dataset. ClueGO and goana tend to yield enrichment results for data with no biological context (i.e., false positives).</p>
</caption>
<graphic xlink:href="fbinf-06-1755664-g002.tif">
<alt-text content-type="machine-generated">Violin plots comparing the number of enriched Gene Ontology terms with adjusted p-value less than 0.05 across different tools: BiNGO, ClueGO, DAVID, Enrichr, GOstats, PANTHER, ShinyGO, WebGestalt, clusterProfiler, gProfiler, goana, and topGO. The plots are shown for input sizes of 50, 100, 200, and 500. Each plot illustrates variations in the number of terms detected by each tool as input size increases.</alt-text>
</graphic>
</fig>
<p>These results reiterate that, while ORA-based methods can produce significant noise, this can be mitigated by appropriately correcting p-values. Likewise, it reinforces the need to test the coherence and filtering capabilities of biological information provided by ORA-based tools.</p>
</sec>
<sec id="s3-2">
<label>3.2</label>
<title>The number of enriched ontologies varies greatly among different tools</title>
<p>One of the main differences we observed when comparing results across enrichment tools was that the number of enriched GO terms varied significantly in the <italic>Hallmark</italic>, <italic>GOBP</italic>, and <italic>Contextual</italic> datasets (<xref ref-type="fig" rid="F3">Figure 3</xref>). DAVID and topGO were the most conservative tools, retrieving considerably fewer enriched categories than the other tools throughout the three datasets. This could be a double-edged sword: while it identifies fewer enriched categories, potentially highlighting the ontologies closely related to the input, it may also overlook unexpected ontologies that could contribute important biological novelty.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Number of enriched ontologies per tool for each input size in the <italic>Hallmark</italic>, <italic>GOBP</italic>, and <italic>Contextual</italic> datasets. In orange, the number of enriched ontologies (adjusted p-value &#x003C;0.05). In blue, the number of non-enriched ontologies (adjusted p-value &#x003E;0.05). <bold>(A)</bold> Number of enriched ontologies for the Hallmark datasets. <bold>(B)</bold> Number of enriched ontologies in the GOBP dataset. <bold>(C)</bold> Number of enriched ontologies in the Contextual datasets. The size of the FEA results varies greatly across tools.</p>
</caption>
<graphic xlink:href="fbinf-06-1755664-g003.tif">
<alt-text content-type="machine-generated">Bar charts illustrating the number of enriched and non-enriched results categorized by adjusted p-values above and below 0.05 across three panels (A, B, and C) for various FEA tools. Each panel represents different analyses: HALLMARK, GOBP, and CONTEXTUAL. Blue bars denote p-values greater than 0.05, and orange bars signify p-values less than 0.05. The x-axis shows the results for the 50, 100, 200, and 500 input sizes for each tool, and the y-axis shows the number of ontologies.</alt-text>
</graphic>
</fig>
<p>In contrast, ClueGO, GOstats, PANTHER, ShinyGO, WebGestalt, clusterProfiler, g:Profiler, and goana yielded the largest results across the analyzed datasets. In this case, identifying too many enriched GO terms may be detrimental, as they are more prone to false positives and allow the user to choose from an excessively broad range of statistically significant terms, thereby biasing the interpretation of the results toward what is relevant to them. We argue that in such cases, a more stringent approach should be taken to analyze the enrichment results, such as using smaller p-value cutoffs and filtering by term annotation size and overlapping genes.</p>
<p>These results indicate that not all tools may be suitable for all potential research questions, as they retrieve highly heterogeneous sets of ontologies. In this sense, conservative tools might not yield biological novelty, as they are more likely to show expected results. In contrast, tools that retrieve thousands of GOs will not only eclipse relevant and expected results but also force researchers to browse overly extensive lists in search of useful findings.</p>
</sec>
<sec id="s3-3">
<label>3.3</label>
<title>The level of biological informativeness of FEA varies across tools</title>
<p>To estimate the biological informativeness of the FEA results provided by each tool, we used the median and average of both the annotation sizes and depths of the enriched GO terms as measures of biological specificity (<xref ref-type="table" rid="T2">Tables 2</xref>, <xref ref-type="table" rid="T3">3</xref>). We chose to limit our analysis to the top 20 and top 100 enriched ontologies, arranged by p-values in ascending order, to lessen the impact of greatly varying result sizes and to examine the tools&#x2019; ranking preferences. Additionally, since the final step of FEA inherently involves manual interpretation, it is reasonable to consider these portions of the results the most relevant. Both selected metrics were analyzed for the top 20 and top 100 ranked enriched ontologies for all FEA results with the <italic>Hallmark</italic> and <italic>Contextual</italic> datasets.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Table compiling the median annotation size (number of genes associated with a pathway) and the median depth (maximal level in the ontology) of the top 20 and top 100 enriched ontologies (adj P-val &#x2264;0.05) for input lists containing 50 and 500 genes for the HALLMARK dataset.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="3" align="left">Tools</th>
<th rowspan="3" align="left">List size</th>
<th colspan="8" align="center">HALLMARK</th>
</tr>
<tr>
<th colspan="2" align="left">Median annotation size</th>
<th colspan="2" align="left">Average annotation size</th>
<th colspan="2" align="left">Median depth</th>
<th colspan="2" align="left">Average depth</th>
</tr>
<tr>
<th align="left">top20</th>
<th align="left">top100</th>
<th align="left">top20</th>
<th align="left">top100</th>
<th align="left">top20</th>
<th align="left">top100</th>
<th align="left">top20</th>
<th align="left">top100</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">BiNGO</td>
<td align="left">50</td>
<td align="left">242.50</td>
<td align="left">172.00</td>
<td align="left">937.85</td>
<td align="left">727.77</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">3.84</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">200.50</td>
<td align="left">218.00</td>
<td align="left">425.65</td>
<td align="left">751.54</td>
<td align="left">4.50</td>
<td align="left">4.00</td>
<td align="left">4.35</td>
<td align="left">3.65</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">152.00</td>
<td align="left">269.00</td>
<td align="left">336.30</td>
<td align="left">839.74</td>
<td align="left">4.50</td>
<td align="left">4.00</td>
<td align="left">4.25</td>
<td align="left">3.79</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">1417.50</td>
<td align="left">380.50</td>
<td align="left">2088.15</td>
<td align="left">903.71</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.35</td>
<td align="left">3.89</td>
</tr>
<tr>
<td align="left">ClueGO</td>
<td align="left">50</td>
<td align="left">176.00</td>
<td align="left">122.00</td>
<td align="left">206.50</td>
<td align="left">141.09</td>
<td align="left">5.00</td>
<td align="left">6.00</td>
<td align="left">5.05</td>
<td align="left">6.42</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">143.00</td>
<td align="left">208.00</td>
<td align="left">202.85</td>
<td align="left">340.03</td>
<td align="left">8.00</td>
<td align="left">6.00</td>
<td align="left">7.60</td>
<td align="left">6.29</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">143.00</td>
<td align="left">310.00</td>
<td align="left">187.85</td>
<td align="left">765.81</td>
<td align="left">8.00</td>
<td align="left">5.00</td>
<td align="left">7.85</td>
<td align="left">5.84</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">597.50</td>
<td align="left">321.00</td>
<td align="left">2385.75</td>
<td align="left">1269.15</td>
<td align="left">5.50</td>
<td align="left">5.00</td>
<td align="left">5.70</td>
<td align="left">5.89</td>
</tr>
<tr>
<td align="left">DAVID</td>
<td align="left">50</td>
<td align="left">55.50</td>
<td align="left">46.50</td>
<td align="left">163.90</td>
<td align="left">112.38</td>
<td align="left">5.00</td>
<td align="left">5.50</td>
<td align="left">6.15</td>
<td align="left">5.72</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">49.00</td>
<td align="left">48.50</td>
<td align="left">100.65</td>
<td align="left">111.19</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">6.65</td>
<td align="left">5.98</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">49.00</td>
<td align="left">51.00</td>
<td align="left">153.25</td>
<td align="left">125.29</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">6.95</td>
<td align="left">5.96</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">82.00</td>
<td align="left">49.50</td>
<td align="left">177.90</td>
<td align="left">118.44</td>
<td align="left">7.50</td>
<td align="left">7.00</td>
<td align="left">7.70</td>
<td align="left">6.77</td>
</tr>
<tr>
<td align="left">ENRICHR</td>
<td align="left">50</td>
<td align="left">67.00</td>
<td align="left">56.50</td>
<td align="left">213.60</td>
<td align="left">168.55</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">5.85</td>
<td align="left">6.58</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">59.50</td>
<td align="left">44.50</td>
<td align="left">158.60</td>
<td align="left">134.59</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">6.20</td>
<td align="left">6.93</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">31.00</td>
<td align="left">49.00</td>
<td align="left">37.15</td>
<td align="left">182.12</td>
<td align="left">9.00</td>
<td align="left">6.00</td>
<td align="left">9.00</td>
<td align="left">6.73</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">54.00</td>
<td align="left">49.00</td>
<td align="left">104.55</td>
<td align="left">140.33</td>
<td align="left">8.00</td>
<td align="left">7.00</td>
<td align="left">7.90</td>
<td align="left">7.47</td>
</tr>
<tr>
<td align="left">GOstats</td>
<td align="left">50</td>
<td align="left">560.00</td>
<td align="left">341.50</td>
<td align="left">744.50</td>
<td align="left">1116.91</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">3.85</td>
<td align="left">5.09</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">121.00</td>
<td align="left">217.00</td>
<td align="left">184.10</td>
<td align="left">616.74</td>
<td align="left">8.00</td>
<td align="left">6.00</td>
<td align="left">7.45</td>
<td align="left">6.17</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">116.00</td>
<td align="left">256.50</td>
<td align="left">156.05</td>
<td align="left">650.91</td>
<td align="left">8.00</td>
<td align="left">5.50</td>
<td align="left">7.95</td>
<td align="left">6.11</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">406.00</td>
<td align="left">337.00</td>
<td align="left">1816.60</td>
<td align="left">1368.79</td>
<td align="left">7.00</td>
<td align="left">5.00</td>
<td align="left">6.40</td>
<td align="left">5.87</td>
</tr>
<tr>
<td align="left">PANTHER</td>
<td align="left">50</td>
<td align="left">320.00</td>
<td align="left">234.00</td>
<td align="left">687.70</td>
<td align="left">945.94</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">4.40</td>
<td align="left">5.33</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">69.00</td>
<td align="left">170.50</td>
<td align="left">98.90</td>
<td align="left">735.85</td>
<td align="left">8.00</td>
<td align="left">5.50</td>
<td align="left">7.95</td>
<td align="left">6.01</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">69.00</td>
<td align="left">250.50</td>
<td align="left">98.90</td>
<td align="left">1080.93</td>
<td align="left">8.00</td>
<td align="left">5.00</td>
<td align="left">7.95</td>
<td align="left">5.85</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">515.50</td>
<td align="left">399.00</td>
<td align="left">3114.90</td>
<td align="left">1607.08</td>
<td align="left">5.50</td>
<td align="left">5.00</td>
<td align="left">5.60</td>
<td align="left">5.63</td>
</tr>
<tr>
<td align="left">ShinyGO</td>
<td align="left">50</td>
<td align="left">749.50</td>
<td align="left">644.50</td>
<td align="left">1267.30</td>
<td align="left">1128.00</td>
<td align="left">3.50</td>
<td align="left">4.00</td>
<td align="left">3.70</td>
<td align="left">4.17</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">297.00</td>
<td align="left">672.00</td>
<td align="left">743.00</td>
<td align="left">1054.52</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">4.45</td>
<td align="left">4.42</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">156.00</td>
<td align="left">637.50</td>
<td align="left">245.20</td>
<td align="left">1114.78</td>
<td align="left">5.00</td>
<td align="left">4.00</td>
<td align="left">5.45</td>
<td align="left">4.61</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">1172.50</td>
<td align="left">652.50</td>
<td align="left">2094.35</td>
<td align="left">1319.82</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.65</td>
<td align="left">4.59</td>
</tr>
<tr>
<td align="left">WebGestalt</td>
<td align="left">50</td>
<td align="left">324.50</td>
<td align="left">206.50</td>
<td align="left">516.10</td>
<td align="left">424.10</td>
<td align="left">4.00</td>
<td align="left">5.00</td>
<td align="left">4.20</td>
<td align="left">5.67</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">115.50</td>
<td align="left">197.00</td>
<td align="left">172.90</td>
<td align="left">476.56</td>
<td align="left">8.00</td>
<td align="left">5.50</td>
<td align="left">7.45</td>
<td align="left">6.05</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">111.00</td>
<td align="left">247.50</td>
<td align="left">152.80</td>
<td align="left">521.66</td>
<td align="left">8.00</td>
<td align="left">5.50</td>
<td align="left">7.70</td>
<td align="left">5.90</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">348.50</td>
<td align="left">244.00</td>
<td align="left">373.60</td>
<td align="left">495.86</td>
<td align="left">7.00</td>
<td align="left">6.00</td>
<td align="left">6.85</td>
<td align="left">6.10</td>
</tr>
<tr>
<td align="left">clusterProfiler</td>
<td align="left">50</td>
<td align="left">191.50</td>
<td align="left">145.00</td>
<td align="left">200.95</td>
<td align="left">170.61</td>
<td align="left">5.00</td>
<td align="left">5.50</td>
<td align="left">5.10</td>
<td align="left">6.18</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">117.00</td>
<td align="left">129.00</td>
<td align="left">155.30</td>
<td align="left">161.49</td>
<td align="left">8.00</td>
<td align="left">6.00</td>
<td align="left">7.65</td>
<td align="left">6.48</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">117.00</td>
<td align="left">156.00</td>
<td align="left">137.60</td>
<td align="left">179.15</td>
<td align="left">8.00</td>
<td align="left">6.00</td>
<td align="left">8.00</td>
<td align="left">6.64</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">159.00</td>
<td align="left">162.50</td>
<td align="left">206.55</td>
<td align="left">183.64</td>
<td align="left">8.00</td>
<td align="left">7.00</td>
<td align="left">7.45</td>
<td align="left">7.04</td>
</tr>
<tr>
<td align="left">gProfiler</td>
<td align="left">50</td>
<td align="left">813.50</td>
<td align="left">552.50</td>
<td align="left">1692.80</td>
<td align="left">1087.44</td>
<td align="left">3.50</td>
<td align="left">4.00</td>
<td align="left">3.45</td>
<td align="left">4.96</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">126.50</td>
<td align="left">263.50</td>
<td align="left">250.50</td>
<td align="left">893.26</td>
<td align="left">7.50</td>
<td align="left">5.00</td>
<td align="left">7.30</td>
<td align="left">5.70</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">113.00</td>
<td align="left">255.50</td>
<td align="left">172.95</td>
<td align="left">743.64</td>
<td align="left">8.00</td>
<td align="left">5.00</td>
<td align="left">7.50</td>
<td align="left">5.98</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">533.50</td>
<td align="left">466.50</td>
<td align="left">1344.90</td>
<td align="left">1437.14</td>
<td align="left">5.00</td>
<td align="left">5.00</td>
<td align="left">5.30</td>
<td align="left">5.51</td>
</tr>
<tr>
<td align="left">goana</td>
<td align="left">50</td>
<td align="left">2688.50</td>
<td align="left">1268.50</td>
<td align="left">5214.05</td>
<td align="left">2832.59</td>
<td align="left">3.00</td>
<td align="left">3.00</td>
<td align="left">2.90</td>
<td align="left">3.79</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">369.00</td>
<td align="left">633.50</td>
<td align="left">3386.25</td>
<td align="left">2186.26</td>
<td align="left">4.50</td>
<td align="left">4.00</td>
<td align="left">5.45</td>
<td align="left">4.94</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">122.00</td>
<td align="left">411.50</td>
<td align="left">2454.80</td>
<td align="left">1978.42</td>
<td align="left">7.50</td>
<td align="left">4.50</td>
<td align="left">6.65</td>
<td align="left">5.42</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">6075.50</td>
<td align="left">1674.50</td>
<td align="left">6540.25</td>
<td align="left">3172.27</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.10</td>
<td align="left">4.74</td>
</tr>
<tr>
<td align="left">topGO</td>
<td align="left">50</td>
<td align="left">22.50</td>
<td align="left">190.50</td>
<td align="left">247.35</td>
<td align="left">352.82</td>
<td align="left">6.00</td>
<td align="left">5.00</td>
<td align="left">5.90</td>
<td align="left">5.19</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">27.50</td>
<td align="left">107.50</td>
<td align="left">61.95</td>
<td align="left">226.10</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">6.70</td>
<td align="left">5.68</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">64.00</td>
<td align="left">47.50</td>
<td align="left">119.70</td>
<td align="left">175.35</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">6.80</td>
<td align="left">5.97</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">86.00</td>
<td align="left">45.00</td>
<td align="left">308.95</td>
<td align="left">154.66</td>
<td align="left">7.00</td>
<td align="left">7.00</td>
<td align="left">7.40</td>
<td align="left">7.35</td>
</tr>
</tbody>
</table>
</table-wrap>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Table compiling the median annotation size (number of genes associated with a pathway) and the median depth (maximal level in the ontology) of the top 20 and top 100 enriched ontologies (adj P-val &#x2264;0.05) for input lists containing 50 and 500 genes for the CONTEXTUAL dataset.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="3" align="left">Tools</th>
<th rowspan="3" align="left">List size</th>
<th colspan="8" align="center">CONTEXTUAL</th>
</tr>
<tr>
<th colspan="2" align="left">Median annotation size</th>
<th colspan="2" align="left">Average annotation size</th>
<th colspan="2" align="left">Median depth</th>
<th colspan="2" align="left">Average depth</th>
</tr>
<tr>
<th align="left">top20</th>
<th align="left">top100</th>
<th align="left">top20</th>
<th align="left">top100</th>
<th align="left">top20</th>
<th align="left">top100</th>
<th align="left">top20</th>
<th align="left">top100</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">BiNGO</td>
<td align="left">50</td>
<td align="left">362.00</td>
<td align="left">238.00</td>
<td align="left">1033.75</td>
<td align="left">518.78</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.00</td>
<td align="left">3.77</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">407.50</td>
<td align="left">159.50</td>
<td align="left">1081.00</td>
<td align="left">439.40</td>
<td align="left">3.50</td>
<td align="left">4.00</td>
<td align="left">4.05</td>
<td align="left">4.50</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">435.00</td>
<td align="left">165.50</td>
<td align="left">901.35</td>
<td align="left">439.74</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.10</td>
<td align="left">4.42</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">503.50</td>
<td align="left">156.50</td>
<td align="left">1147.35</td>
<td align="left">505.01</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">2.90</td>
<td align="left">4.21</td>
</tr>
<tr>
<td align="left">ClueGO</td>
<td align="left">50</td>
<td align="left">81.50</td>
<td align="left">82.00</td>
<td align="left">89.55</td>
<td align="left">84.78</td>
<td align="left">5.50</td>
<td align="left">6.00</td>
<td align="left">5.45</td>
<td align="left">6.12</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">41.00</td>
<td align="left">70.50</td>
<td align="left">171.70</td>
<td align="left">128.64</td>
<td align="left">5.50</td>
<td align="left">5.50</td>
<td align="left">5.85</td>
<td align="left">5.81</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">452.00</td>
<td align="left">287.00</td>
<td align="left">726.60</td>
<td align="left">502.23</td>
<td align="left">3.50</td>
<td align="left">5.00</td>
<td align="left">4.15</td>
<td align="left">5.27</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">386.00</td>
<td align="left">413.50</td>
<td align="left">1238.75</td>
<td align="left">1237.82</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">4.30</td>
<td align="left">4.80</td>
</tr>
<tr>
<td align="left">DAVID</td>
<td align="left">50</td>
<td align="left">52.50</td>
<td align="left">121.00</td>
<td align="left">142.85</td>
<td align="left">267.95</td>
<td align="left">5.50</td>
<td align="left">6.00</td>
<td align="left">5.40</td>
<td align="left">5.82</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">57.50</td>
<td align="left">48.50</td>
<td align="left">118.50</td>
<td align="left">109.42</td>
<td align="left">5.00</td>
<td align="left">5.50</td>
<td align="left">5.20</td>
<td align="left">5.83</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">62.00</td>
<td align="left">52.50</td>
<td align="left">126.00</td>
<td align="left">136.57</td>
<td align="left">5.50</td>
<td align="left">6.00</td>
<td align="left">5.55</td>
<td align="left">5.85</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">63.00</td>
<td align="left">48.00</td>
<td align="left">103.90</td>
<td align="left">98.56</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">5.75</td>
<td align="left">5.59</td>
</tr>
<tr>
<td align="left">ENRICHR</td>
<td align="left">50</td>
<td align="left">53.50</td>
<td align="left">107.50</td>
<td align="left">100.60</td>
<td align="left">191.32</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">5.95</td>
<td align="left">6.14</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">29.00</td>
<td align="left">44.00</td>
<td align="left">62.75</td>
<td align="left">100.97</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">6.20</td>
<td align="left">6.59</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">50.50</td>
<td align="left">35.50</td>
<td align="left">121.25</td>
<td align="left">125.01</td>
<td align="left">5.50</td>
<td align="left">6.00</td>
<td align="left">6.35</td>
<td align="left">6.61</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">74.00</td>
<td align="left">54.50</td>
<td align="left">123.65</td>
<td align="left">137.87</td>
<td align="left">5.50</td>
<td align="left">6.00</td>
<td align="left">6.25</td>
<td align="left">6.40</td>
</tr>
<tr>
<td align="left">GOstats</td>
<td align="left">50</td>
<td align="left">815.50</td>
<td align="left">266.50</td>
<td align="left">1351.05</td>
<td align="left">985.37</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.35</td>
<td align="left">4.64</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">511.50</td>
<td align="left">350.50</td>
<td align="left">1493.20</td>
<td align="left">1048.72</td>
<td align="left">4.50</td>
<td align="left">4.00</td>
<td align="left">4.05</td>
<td align="left">4.63</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">324.00</td>
<td align="left">323.50</td>
<td align="left">1145.40</td>
<td align="left">932.89</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">4.15</td>
<td align="left">4.83</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">371.50</td>
<td align="left">388.00</td>
<td align="left">1176.10</td>
<td align="left">1106.70</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">4.35</td>
<td align="left">4.86</td>
</tr>
<tr>
<td align="left">PANTHER</td>
<td align="left">50</td>
<td align="left">629.00</td>
<td align="left">172.00</td>
<td align="left">1395.40</td>
<td align="left">1044.67</td>
<td align="left">3.50</td>
<td align="left">5.00</td>
<td align="left">3.80</td>
<td align="left">4.86</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">287.00</td>
<td align="left">270.50</td>
<td align="left">1404.80</td>
<td align="left">845.76</td>
<td align="left">5.00</td>
<td align="left">5.00</td>
<td align="left">4.90</td>
<td align="left">4.97</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">843.00</td>
<td align="left">271.00</td>
<td align="left">1631.10</td>
<td align="left">974.72</td>
<td align="left">3.00</td>
<td align="left">5.00</td>
<td align="left">3.25</td>
<td align="left">5.19</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">352.00</td>
<td align="left">320.00</td>
<td align="left">940.85</td>
<td align="left">1369.69</td>
<td align="left">4.00</td>
<td align="left">4.50</td>
<td align="left">4.05</td>
<td align="left">5.08</td>
</tr>
<tr>
<td align="left">ShinyGO</td>
<td align="left">50</td>
<td align="left">638.50</td>
<td align="left">334.00</td>
<td align="left">842.30</td>
<td align="left">838.00</td>
<td align="left">4.00</td>
<td align="left">5.00</td>
<td align="left">4.40</td>
<td align="left">4.67</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">357.50</td>
<td align="left">334.00</td>
<td align="left">821.70</td>
<td align="left">746.90</td>
<td align="left">5.00</td>
<td align="left">4.00</td>
<td align="left">4.50</td>
<td align="left">4.70</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">377.00</td>
<td align="left">426.00</td>
<td align="left">844.45</td>
<td align="left">780.82</td>
<td align="left">4.00</td>
<td align="left">5.00</td>
<td align="left">4.15</td>
<td align="left">5.19</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">355.50</td>
<td align="left">378.50</td>
<td align="left">684.85</td>
<td align="left">882.43</td>
<td align="left">4.50</td>
<td align="left">4.00</td>
<td align="left">4.95</td>
<td align="left">4.92</td>
</tr>
<tr>
<td align="left">WebGestalt</td>
<td align="left">50</td>
<td align="left">562.00</td>
<td align="left">189.00</td>
<td align="left">603.80</td>
<td align="left">485.86</td>
<td align="left">4.00</td>
<td align="left">5.00</td>
<td align="left">3.90</td>
<td align="left">5.15</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">318.00</td>
<td align="left">296.50</td>
<td align="left">657.95</td>
<td align="left">481.40</td>
<td align="left">4.50</td>
<td align="left">5.00</td>
<td align="left">4.20</td>
<td align="left">4.84</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">318.00</td>
<td align="left">277.50</td>
<td align="left">644.05</td>
<td align="left">476.15</td>
<td align="left">4.00</td>
<td align="left">4.00</td>
<td align="left">4.35</td>
<td align="left">4.95</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">317.50</td>
<td align="left">306.50</td>
<td align="left">568.90</td>
<td align="left">496.28</td>
<td align="left">5.00</td>
<td align="left">5.00</td>
<td align="left">4.95</td>
<td align="left">5.21</td>
</tr>
<tr>
<td align="left">clusterProfiler</td>
<td align="left">50</td>
<td align="left">114.00</td>
<td align="left">100.00</td>
<td align="left">154.05</td>
<td align="left">155.87</td>
<td align="left">5.00</td>
<td align="left">6.00</td>
<td align="left">5.60</td>
<td align="left">6.08</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">174.50</td>
<td align="left">67.50</td>
<td align="left">187.25</td>
<td align="left">137.12</td>
<td align="left">5.00</td>
<td align="left">6.00</td>
<td align="left">5.50</td>
<td align="left">5.73</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">170.00</td>
<td align="left">96.50</td>
<td align="left">211.65</td>
<td align="left">153.25</td>
<td align="left">5.50</td>
<td align="left">6.00</td>
<td align="left">5.60</td>
<td align="left">6.10</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">170.00</td>
<td align="left">135.50</td>
<td align="left">216.35</td>
<td align="left">180.17</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">6.05</td>
<td align="left">6.07</td>
</tr>
<tr>
<td align="left">gProfiler</td>
<td align="left">50</td>
<td align="left">1321.00</td>
<td align="left">623.00</td>
<td align="left">3375.20</td>
<td align="left">1966.80</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">2.70</td>
<td align="left">4.05</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">1344.50</td>
<td align="left">697.50</td>
<td align="left">2385.50</td>
<td align="left">1520.07</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.15</td>
<td align="left">4.22</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">1544.00</td>
<td align="left">711.00</td>
<td align="left">3040.55</td>
<td align="left">1861.28</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.00</td>
<td align="left">3.97</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">2442.50</td>
<td align="left">726.00</td>
<td align="left">4509.15</td>
<td align="left">2094.63</td>
<td align="left">2.50</td>
<td align="left">4.00</td>
<td align="left">2.85</td>
<td align="left">4.46</td>
</tr>
<tr>
<td align="left">goana</td>
<td align="left">50</td>
<td align="left">2966.00</td>
<td align="left">655.50</td>
<td align="left">4182.35</td>
<td align="left">2089.35</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">2.55</td>
<td align="left">4.17</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">1491.00</td>
<td align="left">898.00</td>
<td align="left">3215.40</td>
<td align="left">2102.89</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">3.45</td>
<td align="left">3.96</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">2823.50</td>
<td align="left">856.00</td>
<td align="left">4270.35</td>
<td align="left">2003.36</td>
<td align="left">2.50</td>
<td align="left">3.50</td>
<td align="left">2.80</td>
<td align="left">3.94</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">3967.00</td>
<td align="left">720.00</td>
<td align="left">4861.75</td>
<td align="left">2006.60</td>
<td align="left">3.00</td>
<td align="left">4.00</td>
<td align="left">2.95</td>
<td align="left">4.49</td>
</tr>
<tr>
<td align="left">topGO</td>
<td align="left">50</td>
<td align="left">394.50</td>
<td align="left">441.50</td>
<td align="left">437.10</td>
<td align="left">793.67</td>
<td align="left">5.00</td>
<td align="left">5.00</td>
<td align="left">5.30</td>
<td align="left">5.27</td>
</tr>
<tr>
<td align="left"/>
<td align="left">100</td>
<td align="left">80.50</td>
<td align="left">151.50</td>
<td align="left">270.75</td>
<td align="left">318.02</td>
<td align="left">7.00</td>
<td align="left">5.00</td>
<td align="left">7.10</td>
<td align="left">5.76</td>
</tr>
<tr>
<td align="left"/>
<td align="left">200</td>
<td align="left">86.00</td>
<td align="left">129.50</td>
<td align="left">251.65</td>
<td align="left">235.31</td>
<td align="left">7.00</td>
<td align="left">6.00</td>
<td align="left">7.20</td>
<td align="left">6.27</td>
</tr>
<tr>
<td align="left"/>
<td align="left">500</td>
<td align="left">60.50</td>
<td align="left">24.50</td>
<td align="left">101.65</td>
<td align="left">111.66</td>
<td align="left">6.00</td>
<td align="left">6.00</td>
<td align="left">6.65</td>
<td align="left">5.95</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>Remarkably, DAVID, Enrichr, clusterProfiler, and topGO displayed the smallest annotation sizes and the highest depths across all list sizes in both datasets (<xref ref-type="table" rid="T2">Tables 2</xref>, <xref ref-type="table" rid="T3">3</xref>), indicating that these tools tend to obtain more specific enriched ontologies and provide more descriptive results. For example, their results with the 500-hallmark list include, among the top-ranking enriched GO terms, &#x201c;nucleotide excision repair&#x201d; (DAVID and Enrichr), &#x201c;purine ribonucleotide catabolic process&#x201d; (clusterProfiler), and &#x201c;canonical glycolysis&#x201d; (topGO), which are more informative compared to the broader ontologies &#x201c;primary metabolic process&#x201d; and &#x201c;response to stimulus&#x201d; (<xref ref-type="sec" rid="s11">Supplementary Table S2</xref>).</p>
<p>Moreover, ClueGO&#x2019;s output contained highly informative ontologies, especially for smaller list sizes (200, 100, and 50 for the Contextual dataset; 50 and 100 for the <italic>Hallmark</italic> dataset), but failed to do so with larger input lists. In this sense, as we tested the 500 lists, generic terms such as &#x201c;metabolic process&#x201d; and &#x201c;biosynthetic process&#x201d; became much more prevalent.</p>
<p>Conversely, g:Profiler and goana produced substantially less informative results, as they retrieved the largest and shallowest GO terms, which tend to be less precise in terms of biological significance. For instance, goana obtained the ontologies &#x201c;regulation of biological process&#x201d; and &#x201c;metabolic process&#x201d;, among multiple other generic descriptions, across all results with the <italic>Contextual</italic> and <italic>Hallmark</italic> lists, respectively. On the other hand, g:Profiler yielded more informative terms than goana, but still included rather broad GO terms, especially with the <italic>Contextual</italic> dataset, such as &#x201c;response to stimulus&#x201d; and &#x201c;biological regulation&#x201d;. However, we observed that both tools obtained smaller, deeper ontologies with the 200- and 100-lists in the <italic>Hallmark</italic> dataset, particularly in the top 20, compared to the top 100.</p>
<p>The remaining tools BiNGO, GOstats, PANTHER, ShinyGO, and WebGestalt provided intermediate results in terms of biological specificity with the Contextual lists. Nonetheless, a pattern similar to g:Profiler and goana with the <italic>Hallmark</italic> lists of sizes 200 and 100 can be observed. For all of these tools, these particular inputs appear to yield more biologically informative results and place them among the highest-ranked terms. Given that the 500 <italic>Hallmark</italic> list is relatively more heterogeneous, as it includes genes from three distinct biological hallmarks (DNA repair, Hypoxia, and Response to UV radiation), while the smaller lists are comprised solely of genes from the hypoxia hallmark, our results suggest that these tools tend to lose specificity when analyzing more heterogeneous gene lists. Furthermore, loss in specificity observed with the 50 <italic>Hallmark</italic> list &#x2013; which also comprises only hypoxia-related genes &#x2013; is likely related to the small size of the input. This reinforces our approach by showing that the nature and size of the input data will strongly affect the quality of the FEA results.</p>
<p>Furthermore, by analyzing the annotation size distributions for the top 20, top 100, and all enriched GO terms from each tool, we found that differences in distribution patterns are much more pronounced in the top 20 and top 100 than across all results. Specifically, DAVID, Enrichr, clusterProfiler, and topGO consistently identify smaller ontologies among their top results, suggesting these tools favor more specific GO terms at the top ranks. In contrast, others spread such precise ontologies throughout their results (<xref ref-type="fig" rid="F4">Figures 4</xref>, <xref ref-type="fig" rid="F5">5</xref>). However, it is essential to note that the output size can influence the comparisons with the full extent of the results, especially in tools that yield larger results. The distribution of GO terms&#x2019; annotation size is skewed, with most ontologies having fewer genes annotated to them (<xref ref-type="bibr" rid="B20">Jelier et al., 2008</xref>). Consequently, as output size increases, the distribution of annotation sizes in FEA results typically shifts toward smaller values. Also, in the case of ShinyGO, which only outputs up to 1,000 enriched GO terms, distributions can be affected, as not all of the enriched ontologies are accounted for.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Distribution of the annotation sizes of the enriched ontologies for each <italic>Hallmark</italic> dataset input in the top 20 (first column), top 100 (second column), and all results (third column). Tools display varying distributions of annotation sizes, with some exhibiting a preference for smaller GO terms. The differences are more prominent within the top-ranking ontologies. <bold>(A&#x2013;C)</bold> Distribution for the 500 Hallmark list in the top 20, 100, and all results, respectively. <bold>(D&#x2013;F)</bold> Distribution for the 200 <italic>Hallmark</italic> list in the top 20, 100, and all results, respectively. <bold>(G&#x2013;I)</bold> Distribution for the 100 <italic>Hallmark</italic> list in the top 20, 100, and all results, respectively. <bold>(J&#x2013;L)</bold> Distribution for the 200 <italic>Hallmark</italic> list in the top 20, 100, and all results, respectively.</p>
</caption>
<graphic xlink:href="fbinf-06-1755664-g004.tif">
<alt-text content-type="machine-generated">Twelve grouped box plots labeled A to L compare the distribution of the annotation sizes of the enrichment results for various tools, including BiNGO, clusterProfiler, DAVID, among others, in the Hallmark dataset. Each plot shows the distribution and variation within each tool in the top-20, top-100, and in all results. Categories are color-coded for visual distinction.</alt-text>
</graphic>
</fig>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Distribution of the annotation sizes of the enriched ontologies for each <italic>Contextual</italic> dataset input in the top 20 (first column), top 100 (second column), and all results (third column). Tools display varying distributions of annotation sizes, with some exhibiting a preference for smaller GO terms. The differences are more prominent within the top-ranking ontologies. <bold>(A&#x2013;C)</bold> Distribution for the 500 <italic>Contextual</italic> list in the top 20, 100, and all results, respectively. <bold>(D&#x2013;F)</bold> Distribution for the 200 <italic>Contextual</italic> list in the top 20, 100, and all results, respectively. <bold>(G&#x2013;I)</bold> Distribution for the 100 <italic>Contextual</italic> list in the top 20, 100, and all results, respectively. <bold>(J&#x2013;L)</bold> Distribution for the 200 <italic>Contextual</italic> list in the top 20, 100, and all results, respectively.</p>
</caption>
<graphic xlink:href="fbinf-06-1755664-g005.tif">
<alt-text content-type="machine-generated">Twelve grouped box plots labeled A to L compare the distribution of the annotation sizes of the enrichment results for various tools, including BiNGO, clusterProfiler, DAVID, among others, in the Contextual dataset. Each plot shows the distribution and variation within each tool in the top-20, top-100, and in all results. Categories are color-coded for visual distinction.</alt-text>
</graphic>
</fig>
<p>The findings above reflect a known issue in the field, widely discussed among users, regarding the relevance of the retrieved GO terms. Users are familiar with the challenge of finding relevant or more descriptive GO terms amidst the vast number of bioprocesses that tools might yield. It is unfeasible or impractical to manually analyze, for instance, more than 5000 GO terms (<xref ref-type="fig" rid="F2">Figure 2</xref>) to identify relevant bioprocesses. Our results indicate that tools differ significantly in the GO terms they rank first and in their level of informativeness.</p>
</sec>
<sec id="s3-4">
<label>3.4</label>
<title>The degree of stringency differs among tools</title>
<p>We also assessed the tools&#x2019; performance in terms of accuracy and FPR by using Wang&#x2019;s SS to construct a confusion matrix for each FEA result with the <italic>GOBP</italic> dataset. DAVID, WebGestalt, and topGO were the tools that exhibited the highest accuracy values (<xref ref-type="table" rid="T4">Table 4</xref>). Since accuracy, in the context of this study and of how the confusion matrix was defined, fundamentally denotes how similar and dissimilar the enriched and not enriched GO terms are to the targets for each list, this indicates that the FEA results produced by the selected tools are highly consistent with the biological context of the input list.</p>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Tools performance metrics for all list sizes in the GOBP dataset.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Tool</th>
<th align="center">Size</th>
<th align="center">Accuracy</th>
<th align="center">FPR</th>
<th align="center">Median target rank</th>
<th align="center">Identified targets</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left" rowspan="4">BiNGO</td>
<td align="right">500</td>
<td align="center">0.43</td>
<td align="center">0.59</td>
<td align="center">220</td>
<td align="center">10/15</td>
</tr>
<tr>
<td align="right">200</td>
<td align="center">0.44</td>
<td align="center">0.57</td>
<td align="center">91.5</td>
<td align="center">4/7</td>
</tr>
<tr>
<td align="right">100</td>
<td align="center">0.76</td>
<td align="center">0.25</td>
<td align="center">5</td>
<td align="center">3/4</td>
</tr>
<tr>
<td align="right">50</td>
<td align="center">0.41</td>
<td align="center">0.59</td>
<td align="center">4</td>
<td align="center">2/3</td>
</tr>
<tr>
<td align="left" rowspan="4">ClueGO</td>
<td align="right">500</td>
<td align="center">0.36</td>
<td align="center">0.66</td>
<td align="center">155</td>
<td align="center">15/15</td>
</tr>
<tr>
<td align="right">200</td>
<td align="center">0.28</td>
<td align="center">0.74</td>
<td align="center">28</td>
<td align="center">7/7</td>
</tr>
<tr>
<td align="right">100</td>
<td align="center">0.16</td>
<td align="center">0.86</td>
<td align="center">3</td>
<td align="center">4/4</td>
</tr>
<tr>
<td align="right">50</td>
<td align="center">0.04</td>
<td align="center">1.00</td>
<td align="center">6</td>
<td align="center">3/3</td>
</tr>
<tr>
<td align="left" rowspan="4">DAVID</td>
<td align="right">500</td>
<td align="center">0.77</td>
<td align="center">0.22</td>
<td align="center">22</td>
<td align="center">14/15</td>
</tr>
<tr>
<td align="right">200</td>
<td align="center">0.81</td>
<td align="center">0.18</td>
<td align="center">17</td>
<td align="center">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.87</td>
<td align="left">0.12</td>
<td align="left">2.5</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.87</td>
<td align="left">0.13</td>
<td align="left">9</td>
<td align="left">2/3</td>
</tr>
<tr>
<td align="left" rowspan="4">Enrichr</td>
<td align="left">500</td>
<td align="left">0.53</td>
<td align="left">0.49</td>
<td align="left">24</td>
<td align="left">13/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.50</td>
<td align="left">0.51</td>
<td align="left">8</td>
<td align="left">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.67</td>
<td align="left">0.34</td>
<td align="left">3</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.48</td>
<td align="left">0.52</td>
<td align="left">1.5</td>
<td align="left">2/3</td>
</tr>
<tr>
<td align="left" rowspan="4">GOstats</td>
<td align="left">500</td>
<td align="left">0.40</td>
<td align="left">0.62</td>
<td align="left">134</td>
<td align="left">15/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.43</td>
<td align="left">0.58</td>
<td align="left">28</td>
<td align="left">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.64</td>
<td align="left">0.37</td>
<td align="left">4</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.44</td>
<td align="left">0.57</td>
<td align="left">6</td>
<td align="left">3/3</td>
</tr>
<tr>
<td align="left" rowspan="4">PANTHER</td>
<td align="left">500</td>
<td align="left">0.45</td>
<td align="left">0.57</td>
<td align="left">163</td>
<td align="left">15/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.49</td>
<td align="left">0.52</td>
<td align="left">37</td>
<td align="left">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.78</td>
<td align="left">0.22</td>
<td align="left">3.5</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.54</td>
<td align="left">0.46</td>
<td align="left">28.5</td>
<td align="left">2/3</td>
</tr>
<tr>
<td align="left" rowspan="4">ShinyGO</td>
<td align="left">500</td>
<td align="left">0.12</td>
<td align="left">1.00</td>
<td align="left">165</td>
<td align="left">15/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.05</td>
<td align="left">1.00</td>
<td align="left">51</td>
<td align="left">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.19</td>
<td align="left">0.82</td>
<td align="left">4</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.02</td>
<td align="left">1.00</td>
<td align="left">4</td>
<td align="left">3/3</td>
</tr>
<tr>
<td align="left" rowspan="4">WebGestalt</td>
<td align="left">500</td>
<td align="left">0.63</td>
<td align="left">0.38</td>
<td align="left">102</td>
<td align="left">15/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.70</td>
<td align="left">0.29</td>
<td align="left">22</td>
<td align="left">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.89</td>
<td align="left">0.11</td>
<td align="left">4.5</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.77</td>
<td align="left">0.23</td>
<td align="left">7</td>
<td align="left">3/3</td>
</tr>
<tr>
<td align="left" rowspan="4">clusterProfiler</td>
<td align="left">500</td>
<td align="left">0.40</td>
<td align="left">0.62</td>
<td align="left">51</td>
<td align="left">15/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.46</td>
<td align="left">0.55</td>
<td align="left">13</td>
<td align="left">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.69</td>
<td align="left">0.31</td>
<td align="left">3</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.41</td>
<td align="left">0.60</td>
<td align="left">6</td>
<td align="left">3/3</td>
</tr>
<tr>
<td align="left" rowspan="4">gProfiler</td>
<td align="left">500</td>
<td align="left">0.33</td>
<td align="left">0.70</td>
<td align="left">168</td>
<td align="left">15/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.36</td>
<td align="left">0.66</td>
<td align="left">43</td>
<td align="left">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.52</td>
<td align="left">0.48</td>
<td align="left">4</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.38</td>
<td align="left">0.63</td>
<td align="left">6</td>
<td align="left">3/3</td>
</tr>
<tr>
<td align="left" rowspan="4">goana</td>
<td align="left">500</td>
<td align="left">0.34</td>
<td align="left">0.69</td>
<td align="left">232</td>
<td align="left">15/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.41</td>
<td align="left">0.60</td>
<td align="left">59</td>
<td align="left">7/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.61</td>
<td align="left">0.39</td>
<td align="left">4</td>
<td align="left">4/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.45</td>
<td align="left">0.56</td>
<td align="left">6</td>
<td align="left">3/3</td>
</tr>
<tr>
<td align="left" rowspan="4">topGO</td>
<td align="left">500</td>
<td align="left">0.79</td>
<td align="left">0.19</td>
<td align="left">17.5</td>
<td align="left">10/15</td>
</tr>
<tr>
<td align="left">200</td>
<td align="left">0.79</td>
<td align="left">0.21</td>
<td align="left">9</td>
<td align="left">6/7</td>
</tr>
<tr>
<td align="left">100</td>
<td align="left">0.74</td>
<td align="left">0.24</td>
<td align="left">1</td>
<td align="left">1/4</td>
</tr>
<tr>
<td align="left">50</td>
<td align="left">0.76</td>
<td align="left">0.24</td>
<td align="left">12</td>
<td align="left">3/3</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>Additionally, these tools displayed the lowest FPR values, indicating that the GO terms they enrich are highly related to the input data. ClueGO and ShinyGO were the least accurate software in this analysis (<xref ref-type="table" rid="T4">Table 4</xref>), primarily due to the generation of numerous false positives (i.e., enriched ontologies unrelated to the context of the input list) (<xref ref-type="table" rid="T4">Table 4</xref>). Moreover, both tools underperformed in terms of FPR due to the nature of their results, which contained few or none non-statistically significant GO terms, thereby biasing FPR towards 1. This result highlights that performance metrics are highly dependent on the method used to define them and the specificities of the tools. Therefore, despite being undoubtedly relevant performance measures, solely relying on such metrics to compare different tools can be problematic in the case of FEA.</p>
<p>To complement the confusion matrix analysis, we evaluated the identification and ranking of the target ontologies for each list size in the <italic>GOBP</italic> dataset. Regarding identification abilities, most tools were able to enrich 90% or more of the proposed target (<xref ref-type="table" rid="T4">Table 4</xref>), with the only exceptions being BiNGO and topGO, which failed to do so (66% and 69% identification rates, respectively). In terms of ranking capabilities, across the 100 and 50 lists, all software ranked the identified target GO terms near the top (<xref ref-type="fig" rid="F6">Figure 6</xref>; <xref ref-type="table" rid="T4">Table 4</xref>). However, most tools lost this ability and increased the dispersion of the targets in the ranking. DAVID, Enrichr, clusterProfiler, and topGO were the only tools that could consistently position the targets within the top-ranking ontologies across all list sizes (<xref ref-type="fig" rid="F6">Figure 6</xref>), suggesting that they tend to prioritize results closely related to the input.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Ranks of the target ontologies for each list in the GOBP dataset. The results show that tools tend to spread the targets throughout the results, especially with bigger inputs. DAVID, Enricher, clusterProfiler, and topGO were less likely to disperse them.</p>
</caption>
<graphic xlink:href="fbinf-06-1755664-g006.tif">
<alt-text content-type="machine-generated">Box plot comparing target ontologies' rank across various bioinformatics tools for different list sizes (50, 100, 200, 500). Tools include BiNGO, ClueGO, DAVID, Enrichr, GOstats, PANTHER, ShinyGO, WebGestalt, clusterProfiler, gProfiler, goana, and topGO. Ranks vary, with notable higher medians for BiNGO and goana at larger list sizes. Outliers are marked as circles.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3-5">
<label>3.5</label>
<title>Biological profiles are coherent across tools, but vary in terms of priorization</title>
<p>We further applied MCL to cluster the enriched ontologies for each FEA result with the <italic>Hallmark</italic> and <italic>Contextual</italic> datasets. Headers were manually assigned to the 15 largest clusters to compare the biological information provided by each tool. Notably, with the <italic>Hallmark</italic> lists, most tools identified clusters directly related to the expected biological processes (<xref ref-type="sec" rid="s11">Supplementary Table S3</xref>), namely, DNA repair, hypoxia, and response to UV radiation, since the inputs comprised their respective hallmark genes. The only exceptions were DAVID for the 50-lists and topGO for both 100 and 50-lists, which did not exhibit hypoxia clusters. Both tools did, in fact, enrich at least one hypoxia-related ontology; that enrichment, however, did not cluster with any other enriched GO terms.</p>
<p>As most tools yielded results that were at least partially expected in the <italic>Hallmark</italic> dataset, the most significant differences ultimately come down to the interpretability of their FEA results. Remarkably, DAVID and topGO, two of the tools that had previously produced the most informative results, also displayed the most coherent and precise clusters of GO terms. This observation aligns with the previously reported high accuracy values. In contrast, tools that generated larger and less specific outputs &#x2013; for instance, BiNGO, ClueGO, GOstats, PANTHER, ShinyGO, WebGestalt, g:Profiler, and goana &#x2013; enriched a greater number of terms that were weakly aligned with the biological context of the inputs. Clusters that pertained to &#x201c;nucleotide metabolism&#x201d;, &#x201c;organ and system development&#x201d;, and &#x201c;regulation of cell fate and metabolism&#x201d; were the most frequently produced groupings across all lists in this dataset. Although these are somewhat related to the expected biological processes, their GO terms are not as strongly associated as terms such as &#x201c;DNA repair&#x201d;, &#x201c;response to oxygen levels&#x201d;, or &#x201c;response to radiation&#x201d;. The presence of large numbers of these closely related yet unspecific ontologies might, depending on the tools&#x2019; ranking abilities, conceal the most insightful or expected descriptions, thus hindering the interpretability of the results.</p>
<p>Moreover, the behaviors observed in the clustering results for the <italic>Contextual</italic> dataset FEA analyses were consistent with those observed in the <italic>Hallmark</italic> dataset (<xref ref-type="sec" rid="s11">Supplementary Table S4</xref>). DAVID still exhibited the most coherent and well-defined clusters for the 500-list, but did not display clustering results for the smaller lists, as it enriched none or only a few GO terms. Similarly, topGO also did not exhibit any clusters for any of the inputs in this dataset for the same reason as DAVID. Nonetheless, in general, the profiles of the biological groups yielded, and, thereby, the overall biological information contained in the enrichment results, were coherent across all tools.</p>
<p>As a consequence, this suggests that the main differences in the biological information contained within the FEA results across all tools lie mainly in their ability to prioritize what&#x2019;s most relevant in the ranking and to enrich more specific GO terms at the top of the results.</p>
</sec>
</sec>
<sec sec-type="discussion" id="s4">
<label>4</label>
<title>Discussion</title>
<p>FEA has become a crucial step in studies that rely on omics data to drive biological discovery, as it provides a facilitated way to interpret such information. Given its valuable role, several performance studies have been conducted to compare the various tools that employ FEA (<xref ref-type="bibr" rid="B39">Tarca et al., 2013</xref>; <xref ref-type="bibr" rid="B4">Bayerlov&#xe1; et al., 2015</xref>; <xref ref-type="bibr" rid="B28">Lim et al., 2018</xref>; <xref ref-type="bibr" rid="B33">Nguyen et al., 2019</xref>; <xref ref-type="bibr" rid="B47">Zyla et al., 2019</xref>; <xref ref-type="bibr" rid="B17">Geistlinger et al., 2021</xref>; <xref ref-type="bibr" rid="B8">Buzzao et al., 2024</xref>). Nevertheless, such studies overlook the biological informativeness provided by FEA while focusing comparisons on metrics that, although indispensable, fail to accurately describe the performance of FEA tools. As the final step of FEA inevitably involves manual interpretation by the user, it is of utmost importance, in the context of FEA, to also assess the tools&#x2019; biological precision and interpretability.</p>
<p>In this study, we addressed the aforementioned limitation by developing a novel benchmarking strategy centered around the biological significance of the results. We also provide an extensive analysis of 12 widely used ORA-based GO FEA tools.</p>
<p>By exploiting the GO structure and assessing the biological specificity of the enriched GO terms through their annotation size and depth, we identified insightful tendencies commonly overlooked in other studies. Our results show that DAVID, Enrichr, clusterProfiler, and topGO yield more informative results than the other tools, especially compared with goana and g:Profiler, which tend to enrich more generic ontologies. It is essential to note that alternative methods exist for evaluating the specificity of a GO term. For instance, Information Content (IC), which is essentially calculated based on the annotation size of an ontology and its offspring, and the number of offspring GO terms, are two viable options (<xref ref-type="bibr" rid="B30">Louie et al., 2010</xref>; <xref ref-type="bibr" rid="B40">Tomczak et al., 2018</xref>). However, the measures employed here, annotation size and depth, are much more straightforward to grasp for the average FEA user, as they are direct properties of the GO and don&#x2019;t require additional techniques to evaluate them.</p>
<p>Nonetheless, there is currently no gold-standard method for assessing GO term biological specificity, and these metrics display inherent limitations. For instance, the annotation size is directly affected by annotation coverage, which is particularly problematic when dealing with non-model organisms. Additionally, two terms being at the same depth in the ontology does not necessarily imply that both descriptions are equally specific. Despite these limitations, leveraging the GO&#x2019;s properties (i.e., term annotation size and depth) to estimate specificity remains an intuitive and valid strategy (<xref ref-type="bibr" rid="B25">Lewin and Grieve, 2006</xref>; <xref ref-type="bibr" rid="B30">Louie et al., 2010</xref>; <xref ref-type="bibr" rid="B40">Tomczak et al., 2018</xref>). Regardless of the chosen method, we advocate incorporating metrics that reflect the informativeness of results in FEA benchmark studies, such as those described above.</p>
<p>Additionally, our results suggest that the interpretability of each tool&#x2019;s results may be a key factor in the differences between the selected tools. In this regard, we argue that output sizes and ranking abilities are direct measures of the intelligibility of the biological context provided by FEA. Large results can be challenging to interpret, and the most insightful GO terms may be overlooked amid the ocean of enriched ontologies. Likewise, the ability to position the most relevant terms at the top of the results is a quality instrumental in the case of FEA. Tools with such tendencies may benefit from the usage of additional techniques that make their results more interpretable. Redundancy reduction strategies have been developed to facilitate the interpretation of FEA results (<xref ref-type="bibr" rid="B49">Jantzen et al., 2011</xref>; <xref ref-type="bibr" rid="B50">Ozisik et al., 2022</xref>); however, it is crucial to ensure that such approaches don&#x2019;t remove the most specific descriptions while retaining the less informative generic terms. For example, using subsets of the GO that contain no redundant terms, such as the GO slim, may decrease redundancy in FEA results but at the cost of informativeness, as these downsized subsets tend to span broader ontologies. Besides, methods that group the enriched GO terms into functional clusters, as we did in this study using MCL, can be compelling for uncovering the biological profile of the results, further increasing interpretability, especially since functionally related ontologies might be scattered across the full extent of the results.</p>
<p>Furthermore, resources that implement GO term clustering commonly rely on the GO structure or semantic similarity to determine functionally related clusters (<xref ref-type="bibr" rid="B23">Klopfenstein et al., 2018</xref>; <xref ref-type="bibr" rid="B44">Xin et al., 2022</xref>; <xref ref-type="bibr" rid="B45">Xu et al., 2024</xref>). Nonetheless, most platforms of this kind are optimized to work only with their own imbued FEA outputs. Tools capable of universally conducting such interpretability-oriented techniques to FEA outputs &#x2013; regardless of the tools that produced them &#x2013; are still needed in the field.</p>
<p>We also tested the tools with four different input sizes to investigate their impact on the generated GO profile. The analyses we conducted revealed that input size primarily affects output sizes and the tools&#x2019; ranking abilities by dispersing the expected targets for larger sizes. Tools are affected differently: DAVID, Enrichr, clusterProfiler, and topGO, which, curiously, were also the ones to provide more biologically specific results, prioritized the targets by placing them closer to the top of the results, and were less prone to dispersing them. This suggests that these tools tend to position the GO terms strongly related to the input among the top-ranking categories, thereby displaying results that are easier for the user to interpret. For the user, this is relevant because it helps select the most appropriate tool based on the research question. In this sense, hypothesis-oriented studies may take advantage of tools that yield larger, less stringent results, while studies aiming to confirm a biological response will mostly benefit from more conservative tools. Likewise, our results might aid researchers in selecting tools that better fit their input size, as different software exhibit heterogeneous ranking and GO identification performance across input sizes.</p>
<p>The results presented here demonstrate that, even when using the same FEA method, different tools can yield substantially different results. These discrepancies are primarily attributable to (i) variations in GO and annotation versions, (ii) modifications within the statistical approaches themselves, and (iii) the unique algorithms each tool employs for statistical testing. The GO database is regularly updated, with terms being added or removed and annotation sources being revised. Such updates can significantly affect FEA results and change their biological interpretation (<xref ref-type="bibr" rid="B40">Tomczak et al., 2018</xref>). Additionally, the implementation of the ORA method and the statistical correction can vary slightly across tools. For example, DAVID uses a modified version of Fisher&#x2019;s Exact Test, known as the EASE score, which subtracts 01 from the number of input genes mapped to a given pathway, thereby conferring more conservative performance than the original Fisher&#x2019;s Exact Test. Similarly, differences in how statistical correction methods are implemented can also influence results and their biological interpretation (<xref ref-type="bibr" rid="B46">Ziemann et al., 2024</xref>). Finally, the algorithms used to perform multiple statistical tests and generate outputs can produce significant differences in results across tools. For instance, the default algorithm in topGO implements a technique that essentially filters out redundant terms while retaining the ontologies that best characterize the input list, thereby preserving its biological relevance. Likewise, g:Profiler, WebGestalt, clusterProfiler, and ClueGO include built-in options that perform similar functions, but they were outside the scope of this study, as they are optional post-processing steps rather than core components of the algorithms. Users should be mindful of the selected tool&#x2019;s specific features to ensure they obtain the most meaningful results from FEA.</p>
<p>Ultimately, the choice of the appropriate tool is determined primarily by the user&#x2019;s analysis objectives and the nature of the input data. For analyses aimed at investigating novel biological pathways that have not been previously described within a given condition, tools that yield a larger number of enriched terms and more specific terms, such as clusterProfiler and Enrichr, are preferable. Furthermore, these tools and others that retrieve broader terms (i.e., BiNGO, ClueGO, GOstats, PANTHER, WebGestalt, and g:Profiler) are also suitable for highly heterogeneous data, as more stringent software may exclude loosely related yet potentially interesting descriptions. Conversely, more conservative approaches, exemplified by tools like DAVID and topGO, may be more suitable when the primary objective is to validate existing hypotheses and the input data is coherent, such as when the biological condition in question is known to be associated with specific bioprocesses or when the study focuses mainly on genes known to participate in overlapping pathways. Moreover, the user should be aware of the tools&#x2019; characteristics and outputs when selecting software and performing FEA. For instance, tools that yield larger and broader results may require filtering and visualization strategies to facilitate interpretation. Many of the tools selected in this study include such features embedded in the software. For example, ClueGO, clusterProfiler, and DAVID provide functional clustering options. Similarly, ClueGO, ShinyGO, WebGestalt, clusterProfiler, and g:Profiler all provide filtering strategies (e.g., redundancy removal functions, GO term size filtering, and GO term depth filtering) and built-in visualizations. Such features significantly increase the usability of FEA tools. Likewise, users should consider tools&#x2019; characteristics that are not directly related to the FEA itself. As mentioned earlier, FEA tools offer a wide range of additional features that can significantly impact analysis efficiency; consequently, the quality of documentation and tutorials plays a central role in FEA tool selection. Exceptionally, all the tools included in our study have satisfactory documentation on practical usability. However, some of them do not report key information, for example, the exact GO and annotation versions being used to conduct FEA, and that is quite problematic in the FEA field, especially because it hinders reproducibility.</p>
<p>Taken together, our results highlight aspects of the FEA tool&#x2019;s behavior that are typically overlooked and provide relevant information that better guides the selection of FEA software.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The input datasets and scripts can be acquired from <ext-link ext-link-type="uri" xlink:href="https://github.com/LARA-Lab-Aging">https://github.com/LARA-Lab-Aging</ext-link>.</p>
</sec>
<sec sec-type="author-contributions" id="s6">
<title>Author contributions</title>
<p>Fd: Conceptualization, Data curation, Formal Analysis, Investigation, Methodology, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review and editing. FG: Formal Analysis, Validation, Writing &#x2013; review and editing. BF: Conceptualization, Data curation, Formal Analysis, Funding acquisition, Investigation, Methodology, Project administration, Resources, Supervision, Visualization, Writing &#x2013; original draft, Writing &#x2013; review and editing.</p>
</sec>
<ack>
<title>Acknowledgements</title>
<p>We thank CAPES, FAPERGS, and CNPq for financial support.</p>
</ack>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of interest</title>
<p>The author(s) declared that this work was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="s9">
<title>Generative AI statement</title>
<p>The author(s) declared that generative AI was not used in the creation of this manuscript.</p>
<p>Any alternative text (alt text) provided alongside figures in this article has been generated by Frontiers with the support of artificial intelligence and reasonable efforts have been made to ensure accuracy, including review by the authors wherever possible. If you identify any issues, please contact us.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec sec-type="supplementary-material" id="s11">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fbinf.2026.1755664/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fbinf.2026.1755664/full&#x23;supplementary-material</ext-link>
</p>
<supplementary-material xlink:href="Table2.xlsx" id="SM1" mimetype="application/xlsx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table3.xlsx" id="SM2" mimetype="application/xlsx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table4.xlsx" id="SM3" mimetype="application/xlsx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table1.xlsx" id="SM4" mimetype="application/xlsx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Table5.xlsx" id="SM5" mimetype="application/xlsx" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image1.pdf" id="SM6" mimetype="application/pdf" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<fn-group>
<fn fn-type="custom" custom-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/802558/overview">Juw Won Park</ext-link>, University of Louisville, United States</p>
</fn>
<fn fn-type="custom" custom-type="reviewed-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1514220/overview">Jun Zhang</ext-link>, China Pharmaceutical University, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/2638296/overview">Eszter Ari</ext-link>, E&#xf6;tv&#xf6;s Lor&#xe1;nd University, Hungary</p>
</fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Agrawal</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Balc&#x131;</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Hanspers</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Coort</surname>
<given-names>S. L.</given-names>
</name>
<name>
<surname>Martens</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Slenter</surname>
<given-names>D. N.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>WikiPathways 2024: next generation pathway database</article-title>. <source>Nucleic Acids Res.</source> <volume>52</volume>, <fpage>D679</fpage>&#x2013;<lpage>D689</lpage>. <pub-id pub-id-type="doi">10.1093/NAR/GKAD960</pub-id>
<pub-id pub-id-type="pmid">37941138</pub-id>
</mixed-citation>
</ref>
<ref id="B2">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Alexa</surname>
<given-names>R. J.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>topGO: enrichment analysis for gene ontology</article-title>. <pub-id pub-id-type="doi">10.18129/B9.bioc.topGO</pub-id>
</mixed-citation>
</ref>
<ref id="B3">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ashburner</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Ball</surname>
<given-names>C. A.</given-names>
</name>
<name>
<surname>Blake</surname>
<given-names>J. A.</given-names>
</name>
<name>
<surname>Botstein</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Butler</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Cherry</surname>
<given-names>J. M.</given-names>
</name>
<etal/>
</person-group> (<year>2000</year>). <article-title>Gene ontology: tool for the unification of biology</article-title>. <source>Nat. Genet.</source> <volume>25</volume> (<issue>1</issue>), <fpage>25</fpage>&#x2013;<lpage>29</lpage>. <pub-id pub-id-type="doi">10.1038/75556</pub-id>
<pub-id pub-id-type="pmid">10802651</pub-id>
</mixed-citation>
</ref>
<ref id="B4">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bayerlov&#xe1;</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Jung</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Kramer</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Klemm</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Bleckmann</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Bei&#xdf;barth</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>Comparative study on gene set and pathway topology-based enrichment methods</article-title>. <source>BMC Bioinforma.</source> <volume>16</volume>. <pub-id pub-id-type="doi">10.1186/s12859-015-0751-5</pub-id>
<pub-id pub-id-type="pmid">26489510</pub-id>
</mixed-citation>
</ref>
<ref id="B5">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Benjamini</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Hochberg</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>1995</year>). <article-title>Controlling the false discovery rate: a practical and powerful approach to multiple</article-title>.</mixed-citation>
</ref>
<ref id="B6">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bindea</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Mlecnik</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Hackl</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Charoentong</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Tosolini</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Kirilovsky</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2009</year>). <article-title>ClueGO: a cytoscape plug-in to decipher functionally grouped gene ontology and pathway annotation networks</article-title>. <source>Bioinformatics</source> <volume>25</volume>, <fpage>1091</fpage>&#x2013;<lpage>1093</lpage>. <pub-id pub-id-type="doi">10.1093/BIOINFORMATICS/BTP101</pub-id>
<pub-id pub-id-type="pmid">19237447</pub-id>
</mixed-citation>
</ref>
<ref id="B7">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Broh&#xe9;e</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>van Helden</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>Evaluation of clustering algorithms for protein-protein interaction networks</article-title>. <source>BMC Bioinforma.</source> <volume>7</volume>, <fpage>488</fpage>. <pub-id pub-id-type="doi">10.1186/1471-2105-7-488</pub-id>
<pub-id pub-id-type="pmid">17087821</pub-id>
</mixed-citation>
</ref>
<ref id="B8">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Buzzao</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Castresana-Aguirre</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Guala</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Sonnhammer</surname>
<given-names>E. L. L.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Benchmarking enrichment analysis methods with the disease pathway network</article-title>. <source>Brief. Bioinform</source> <volume>25</volume>, <fpage>bbae069</fpage>. <pub-id pub-id-type="doi">10.1093/bib/bbae069</pub-id>
<pub-id pub-id-type="pmid">38436561</pub-id>
</mixed-citation>
</ref>
<ref id="B9">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>E. Y.</given-names>
</name>
<name>
<surname>Tan</surname>
<given-names>C. M.</given-names>
</name>
<name>
<surname>Kou</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Duan</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Meirelles</surname>
<given-names>G. V.</given-names>
</name>
<etal/>
</person-group> (<year>2013</year>). <article-title>Enrichr: interactive and collaborative HTML5 gene list enrichment analysis tool</article-title>. <source>BMC Bioinforma.</source> <volume>14</volume>, <fpage>128</fpage>. <pub-id pub-id-type="doi">10.1186/1471-2105-14-128</pub-id>
<pub-id pub-id-type="pmid">23586463</pub-id>
</mixed-citation>
</ref>
<ref id="B10">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Consortium</surname>
<given-names>T. G. O.</given-names>
</name>
<name>
<surname>Aleksander</surname>
<given-names>S. A.</given-names>
</name>
<name>
<surname>Balhoff</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Carbon</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Cherry</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Drabkin</surname>
<given-names>H. J.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>The gene ontology knowledgebase in 2023</article-title>. <source>Genetics</source> <volume>224</volume>, <fpage>iyad031</fpage>. <pub-id pub-id-type="doi">10.1093/GENETICS/IYAD031</pub-id>
<pub-id pub-id-type="pmid">36866529</pub-id>
</mixed-citation>
</ref>
<ref id="B11">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dong</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Hao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>LEGO: a novel method for gene set over-representation analysis by incorporating network-based gene weights</article-title>. <source>Sci. Rep.</source> <volume>6</volume>, <fpage>18871</fpage>. <pub-id pub-id-type="doi">10.1038/srep18871</pub-id>
<pub-id pub-id-type="pmid">26750448</pub-id>
</mixed-citation>
</ref>
<ref id="B12">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Durinck</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Moreau</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Kasprzyk</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Davis</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>De Moor</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Brazma</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2005</year>). <article-title>BioMart and bioconductor: a powerful link between biological databases and microarray data analysis</article-title>. <source>Bioinformatics</source> <volume>21</volume>, <fpage>3439</fpage>&#x2013;<lpage>3440</lpage>. <pub-id pub-id-type="doi">10.1093/BIOINFORMATICS/BTI525</pub-id>
<pub-id pub-id-type="pmid">16082012</pub-id>
</mixed-citation>
</ref>
<ref id="B13">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Durinck</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Spellman</surname>
<given-names>P. T.</given-names>
</name>
<name>
<surname>Birney</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Huber</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Mapping identifiers for the integration of genomic datasets with the R/Bioconductor package biomaRt</article-title>. <source>Nat. Protoc.</source> <volume>4</volume> (<issue>8</issue>), <fpage>1184</fpage>&#x2013;<lpage>1191</lpage>. <pub-id pub-id-type="doi">10.1038/nprot.2009.97</pub-id>
<pub-id pub-id-type="pmid">19617889</pub-id>
</mixed-citation>
</ref>
<ref id="B14">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Elizarraras</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Liao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Zhu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Pico</surname>
<given-names>A. R.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>WebGestalt 2024: faster gene set analysis and new support for metabolomics and multi-omics</article-title>. <source>Nucleic Acids Res.</source> <volume>52</volume>, <fpage>W415</fpage>&#x2013;<lpage>W421</lpage>. <pub-id pub-id-type="doi">10.1093/NAR/GKAE456</pub-id>
<pub-id pub-id-type="pmid">38808672</pub-id>
</mixed-citation>
</ref>
<ref id="B15">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Falcon</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Gentleman</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>Using GOstats to test gene lists for GO term association</article-title>. <source>Bioinformatics</source> <volume>23</volume>, <fpage>257</fpage>&#x2013;<lpage>258</lpage>. <pub-id pub-id-type="doi">10.1093/BIOINFORMATICS/BTL567</pub-id>
<pub-id pub-id-type="pmid">17098774</pub-id>
</mixed-citation>
</ref>
<ref id="B16">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ge</surname>
<given-names>S. X.</given-names>
</name>
<name>
<surname>Jung</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Jung</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>ShinyGO: a graphical gene-set enrichment tool for animals and plants</article-title>. <source>Bioinformatics</source> <volume>36</volume>, <fpage>2628</fpage>&#x2013;<lpage>2629</lpage>. <pub-id pub-id-type="doi">10.1093/BIOINFORMATICS/BTZ931</pub-id>
<pub-id pub-id-type="pmid">31882993</pub-id>
</mixed-citation>
</ref>
<ref id="B17">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Geistlinger</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Csaba</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Santarelli</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Ramos</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Schiffer</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Turaga</surname>
<given-names>N.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Toward a gold standard for benchmarking gene set enrichment analysis</article-title>. <source>Brief. Bioinform</source> <volume>22</volume>, <fpage>545</fpage>&#x2013;<lpage>556</lpage>. <pub-id pub-id-type="doi">10.1093/bib/bbz158</pub-id>
<pub-id pub-id-type="pmid">32026945</pub-id>
</mixed-citation>
</ref>
<ref id="B18">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>D. W.</given-names>
</name>
<name>
<surname>Sherman</surname>
<given-names>B. T.</given-names>
</name>
<name>
<surname>Lempicki</surname>
<given-names>R. A.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Systematic and integrative analysis of large gene lists using DAVID bioinformatics resources</article-title>. <source>Nat. Protoc.</source> <volume>4</volume>, <fpage>44</fpage>&#x2013;<lpage>57</lpage>. <pub-id pub-id-type="doi">10.1038/NPROT.2008.211</pub-id>
<pub-id pub-id-type="pmid">19131956</pub-id>
</mixed-citation>
</ref>
<ref id="B19">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hung</surname>
<given-names>J. H.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>T. H.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Weng</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>DeLisi</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2012</year>). <article-title>Gene set enrichment analysis: performance evaluation and usage guidelines</article-title>. <source>Brief. Bioinform</source> <volume>13</volume>, <fpage>281</fpage>&#x2013;<lpage>291</lpage>. <pub-id pub-id-type="doi">10.1093/bib/bbr049</pub-id>
<pub-id pub-id-type="pmid">21900207</pub-id>
</mixed-citation>
</ref>
<ref id="B49">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jantzen</surname>
<given-names>S. G.</given-names>
</name>
<name>
<surname>Sutherland</surname>
<given-names>B. J.</given-names>
</name>
<name>
<surname>Minkley</surname>
<given-names>D. R.</given-names>
</name>
<name>
<surname>Koop</surname>
<given-names>B. F.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>GO trimming: systematically reducing redundancy in large Gene Ontology datasets</article-title>. <source>BMC Res. Notes</source> <volume>4</volume>, <fpage>267</fpage> <pub-id pub-id-type="doi">10.1186/1756-0500-4-267</pub-id>
<pub-id pub-id-type="pmid">21798041</pub-id>
</mixed-citation>
</ref>
<ref id="B20">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jelier</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Schuemie</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Roes</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Vanmulligen</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Kors</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Literature-based concept profiles for gene annotation: the issue of weighting</article-title>. <source>Int. J. Med. Inf.</source> <volume>77</volume>, <fpage>354</fpage>&#x2013;<lpage>362</lpage>. <pub-id pub-id-type="doi">10.1016/j.ijmedinf.2007.07.004</pub-id>
<pub-id pub-id-type="pmid">17827057</pub-id>
</mixed-citation>
</ref>
<ref id="B21">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kanehisa</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Goto</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2000</year>). <article-title>KEGG: kyoto encyclopedia of genes and genomes</article-title>. <source>Nucleic Acids Res.</source> <volume>28</volume>, <fpage>27</fpage>&#x2013;<lpage>30</lpage>. <pub-id pub-id-type="doi">10.1093/NAR/28.1.27</pub-id>
<pub-id pub-id-type="pmid">10592173</pub-id>
</mixed-citation>
</ref>
<ref id="B22">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kanehisa</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Furumichi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Sato</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Matsuura</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Ishiguro-Watanabe</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2025</year>). <article-title>KEGG: biological systems database as a model of the real world</article-title>. <source>Nucleic Acids Res.</source> <volume>53</volume>, <fpage>D672</fpage>&#x2013;<lpage>D677</lpage>. <pub-id pub-id-type="doi">10.1093/NAR/GKAE909</pub-id>
<pub-id pub-id-type="pmid">39417505</pub-id>
</mixed-citation>
</ref>
<ref id="B23">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Klopfenstein</surname>
<given-names>D. V.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Pedersen</surname>
<given-names>B. S.</given-names>
</name>
<name>
<surname>Ram&#xed;rez</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Vesztrocy</surname>
<given-names>A. W.</given-names>
</name>
<name>
<surname>Naldi</surname>
<given-names>A.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>GOATOOLS: a python library for gene ontology analyses</article-title>. <source>Sci. Rep.</source> <volume>8</volume>, <fpage>10872</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-018-28948-z</pub-id>
<pub-id pub-id-type="pmid">30022098</pub-id>
</mixed-citation>
</ref>
<ref id="B24">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kolberg</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Raudvere</surname>
<given-names>U.</given-names>
</name>
<name>
<surname>Kuzmin</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Adler</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Vilo</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Peterson</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>g:Profiler&#x2014;interoperable web service for functional enrichment analysis and gene identifier mapping (2023 update)</article-title>. <source>Nucleic Acids Res.</source> <volume>51</volume>, <fpage>W207</fpage>&#x2013;<lpage>W212</lpage>. <pub-id pub-id-type="doi">10.1093/NAR/GKAD347</pub-id>
<pub-id pub-id-type="pmid">37144459</pub-id>
</mixed-citation>
</ref>
<ref id="B25">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lewin</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Grieve</surname>
<given-names>I. C.</given-names>
</name>
</person-group> (<year>2006</year>). <article-title>Grouping gene ontology terms to improve the assessment of gene set enrichment in microarray data</article-title>. <source>BMC Bioinforma.</source> <volume>7</volume>, <fpage>426</fpage>. <pub-id pub-id-type="doi">10.1186/1471-2105-7-426</pub-id>
<pub-id pub-id-type="pmid">17018143</pub-id>
</mixed-citation>
</ref>
<ref id="B26">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liberzon</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Subramanian</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Pinchback</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Thorvaldsd&#xf3;ttir</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Tamayo</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Mesirov</surname>
<given-names>J. P.</given-names>
</name>
</person-group> (<year>2011</year>). <article-title>Molecular signatures database (MSigDB) 3.0</article-title>. <source>Bioinformatics</source> <volume>27</volume>, <fpage>1739</fpage>&#x2013;<lpage>1740</lpage>. <pub-id pub-id-type="doi">10.1093/BIOINFORMATICS/BTR260</pub-id>
<pub-id pub-id-type="pmid">21546393</pub-id>
</mixed-citation>
</ref>
<ref id="B27">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liberzon</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Birger</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Thorvaldsd&#xf3;ttir</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Ghandi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Mesirov</surname>
<given-names>J. P.</given-names>
</name>
<name>
<surname>Tamayo</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2015</year>). <article-title>The molecular signatures database (MSigDB) hallmark gene set collection</article-title>. <source>Cell Syst.</source> <volume>1</volume>, <fpage>417</fpage>&#x2013;<lpage>425</lpage>. <pub-id pub-id-type="doi">10.1016/J.CELS.2015.12.004</pub-id>
<pub-id pub-id-type="pmid">26771021</pub-id>
</mixed-citation>
</ref>
<ref id="B28">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lim</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Jung</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Rhee</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Kim</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Comprehensive and critical evaluation of individualized pathway activity measurement tools on pan-cancer data</article-title>. <source>Brief. Bioinform</source> <volume>21</volume>, <fpage>36</fpage>&#x2013;<lpage>46</lpage>. <pub-id pub-id-type="doi">10.1093/bib/bby097</pub-id>
<pub-id pub-id-type="pmid">30462155</pub-id>
</mixed-citation>
</ref>
<ref id="B29">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lim</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>I.</given-names>
</name>
<name>
<surname>Seo</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Kang</surname>
<given-names>U.</given-names>
</name>
<name>
<surname>Sael</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>PS-MCL: parallel shotgun coarsened Markov clustering of protein interaction networks</article-title>. <source>BMC Bioinforma.</source> <volume>20</volume>, <fpage>381</fpage>. <pub-id pub-id-type="doi">10.1186/s12859-019-2856-8</pub-id>
<pub-id pub-id-type="pmid">31337329</pub-id>
</mixed-citation>
</ref>
<ref id="B30">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Louie</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Bergen</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Higdon</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Kolker</surname>
<given-names>E.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Quantifying protein function specificity in the gene ontology</article-title>. <source>Stand Genomic Sci.</source> <volume>2</volume>, <fpage>238</fpage>&#x2013;<lpage>244</lpage>. <pub-id pub-id-type="doi">10.4056/sigs.561626</pub-id>
<pub-id pub-id-type="pmid">21304708</pub-id>
</mixed-citation>
</ref>
<ref id="B51">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Maere</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Karel</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Martin</surname>
<given-names>K.</given-names>
</name>
</person-group>(<year>2005</year>). <article-title>BiNGO: a cytoscape plugin to assess overrepresentation of gene ontology categories in biological networks</article-title>. <source>Bioinformatics</source> <volume>21</volume>(<issue>16</issue>), <fpage>3448</fpage>&#x2013;<lpage>3449</lpage>.<pub-id pub-id-type="pmid">15972284</pub-id>
</mixed-citation>
</ref>
<ref id="B31">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mi</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Muruganujan</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Ebert</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Mills</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Guo</surname>
<given-names>X.</given-names>
</name>
<etal/>
</person-group> (<year>2019</year>). <article-title>Protocol update for large-scale genome and gene function analysis with the PANTHER classification system (v.14.0)</article-title>. <source>Nat. Protoc.</source> <volume>14</volume> (<issue>14</issue>), <fpage>703</fpage>&#x2013;<lpage>721</lpage>. <pub-id pub-id-type="doi">10.1038/s41596-019-0128-8</pub-id>
<pub-id pub-id-type="pmid">30804569</pub-id>
</mixed-citation>
</ref>
<ref id="B32">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Milacic</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Beavers</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Conley</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Gong</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Gillespie</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Griss</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>The reactome pathway knowledgebase 2024</article-title>. <source>Nucleic Acids Res.</source> <volume>52</volume>, <fpage>D672</fpage>&#x2013;<lpage>D678</lpage>. <pub-id pub-id-type="doi">10.1093/NAR/GKAD1025</pub-id>
<pub-id pub-id-type="pmid">37941124</pub-id>
</mixed-citation>
</ref>
<ref id="B33">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nguyen</surname>
<given-names>T. M.</given-names>
</name>
<name>
<surname>Shafi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Draghici</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Identifying significantly impacted pathways: a comprehensive review and assessment</article-title>. <source>Genome Biol.</source> <volume>20</volume>, <fpage>203</fpage>. <pub-id pub-id-type="doi">10.1186/s13059-019-1790-4</pub-id>
<pub-id pub-id-type="pmid">31597578</pub-id>
</mixed-citation>
</ref>
<ref id="B34">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nunes</surname>
<given-names>I. J. G.</given-names>
</name>
<name>
<surname>Recamonde-Mendoza</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Feltes</surname>
<given-names>B. C.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Gene expression analysis platform (GEAP): a highly customizable, fast, versatile and ready-to-use microarray analysis platform</article-title>. <source>Genet. Mol. Biol.</source> <volume>45</volume>, <fpage>e20210077</fpage>. <pub-id pub-id-type="doi">10.1590/1678-4685-GMB-2021-0077</pub-id>
<pub-id pub-id-type="pmid">34927664</pub-id>
</mixed-citation>
</ref>
<ref id="B50">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ozisik</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>T&#xE9;r&#xE9;zol</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Baudot</surname>
<given-names>A.</given-names>
</name>
</person-group>(<year>2022</year>). <article-title>Orsum: a python package for filtering and comparing enrichment analyses using a simple principle</article-title>. <source>BMC Bioinform.</source> <volume>23</volume>. <pub-id pub-id-type="doi">10.1186/s12859-022-04828-2</pub-id>
<pub-id pub-id-type="pmid">35870894</pub-id>
</mixed-citation>
</ref>
<ref id="B35">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ritchie</surname>
<given-names>M. E.</given-names>
</name>
<name>
<surname>Phipson</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Law</surname>
<given-names>C. W.</given-names>
</name>
<name>
<surname>Shi</surname>
<given-names>W.</given-names>
</name>
<etal/>
</person-group> (<year>2015</year>). <article-title>Limma powers differential expression analyses for RNA-sequencing and microarray studies</article-title>. <source>Nucleic Acids Res.</source> <volume>43</volume>, <fpage>e47</fpage>. <pub-id pub-id-type="doi">10.1093/NAR/GKV007</pub-id>
<pub-id pub-id-type="pmid">25605792</pub-id>
</mixed-citation>
</ref>
<ref id="B48">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sanchez-Palencia</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Gomez-Morales</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Gomez-Capilla</surname>
<given-names>J. A.</given-names>
</name>
<name>
<surname>Pedraza</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Boyero</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Rosell</surname>
<given-names>R.</given-names>
</name>
<etal/>
</person-group> (<year>2011</year>). <article-title>Gene expression profiling reveals novel biomarkers in nonsmall cell lung cancer</article-title>. <source>Int. J. Cancer</source> <volume>129</volume> (<issue>2</issue>), <fpage>355</fpage>&#x2013;<lpage>364</lpage>.<pub-id pub-id-type="pmid">20878980</pub-id>
</mixed-citation>
</ref>
<ref id="B36">
<mixed-citation publication-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Satuluri</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Parthasarathy</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ucar</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2010</year>). &#x201c;<article-title>Markov clustering of protein interaction networks with improved balance and scalability</article-title>,&#x201d; in <conf-name>Proceedings of the First ACM International Conference on Bioinformatics and Computational Biology</conf-name> (<publisher-loc>New York, NY, USA</publisher-loc>: <publisher-name>ACM</publisher-name>), <fpage>247</fpage>&#x2013;<lpage>256</lpage>. <pub-id pub-id-type="doi">10.1145/1854776.1854812</pub-id>
</mixed-citation>
</ref>
<ref id="B37">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sherman</surname>
<given-names>B. T.</given-names>
</name>
<name>
<surname>Hao</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Qiu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jiao</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Baseler</surname>
<given-names>M. W.</given-names>
</name>
<name>
<surname>Lane</surname>
<given-names>H. C.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>DAVID: a web server for functional enrichment analysis and functional annotation of gene lists (2021 update)</article-title>. <source>Nucleic Acids Res.</source> <volume>50</volume>, <fpage>W216</fpage>&#x2013;<lpage>W221</lpage>. <pub-id pub-id-type="doi">10.1093/nar/gkac194</pub-id>
<pub-id pub-id-type="pmid">35325185</pub-id>
</mixed-citation>
</ref>
<ref id="B38">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Subramanian</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Tamayo</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Mootha</surname>
<given-names>V. K.</given-names>
</name>
<name>
<surname>Mukherjee</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ebert</surname>
<given-names>B. L.</given-names>
</name>
<name>
<surname>Gillette</surname>
<given-names>M. A.</given-names>
</name>
<etal/>
</person-group> (<year>2005</year>). <article-title>Gene set enrichment analysis: a knowledge-based approach for interpreting genome-wide expression profiles</article-title>. <source>Proc. Natl. Acad. Sci. U. S. A.</source> <volume>102</volume>, <fpage>15545</fpage>&#x2013;<lpage>15550</lpage>. <pub-id pub-id-type="doi">10.1073/PNAS.0506580102</pub-id>
<pub-id pub-id-type="pmid">16199517</pub-id>
</mixed-citation>
</ref>
<ref id="B39">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tarca</surname>
<given-names>A. L.</given-names>
</name>
<name>
<surname>Bhatti</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Romero</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>A comparison of gene set analysis methods in terms of sensitivity, prioritization and specificity</article-title>. <source>PLoS One</source> <volume>8</volume>, <fpage>e79217</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0079217</pub-id>
<pub-id pub-id-type="pmid">24260172</pub-id>
</mixed-citation>
</ref>
<ref id="B40">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tomczak</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Mortensen</surname>
<given-names>J. M.</given-names>
</name>
<name>
<surname>Winnenburg</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Alessi</surname>
<given-names>D. T.</given-names>
</name>
<name>
<surname>Swamy</surname>
<given-names>V.</given-names>
</name>
<etal/>
</person-group> (<year>2018</year>). <article-title>Interpretation of biological experiments changes with evolution of the gene ontology and its annotations</article-title>. <source>Sci. Rep. 2018</source> <volume>8</volume> (<issue>8</issue>), <fpage>1</fpage>&#x2013;<lpage>10</lpage>. <pub-id pub-id-type="doi">10.1038/s41598-018-23395-2</pub-id>
<pub-id pub-id-type="pmid">29572502</pub-id>
</mixed-citation>
</ref>
<ref id="B41">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Van Dongen</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Graph clustering <italic>via</italic> a discrete uncoupling process</article-title>. <source>SIAM J. Matrix Analysis Appl.</source> <volume>30</volume>, <fpage>121</fpage>&#x2013;<lpage>141</lpage>. <pub-id pub-id-type="doi">10.1137/040608635</pub-id>
</mixed-citation>
</ref>
<ref id="B42">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>J. Z.</given-names>
</name>
<name>
<surname>Du</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Payattakool</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>P. S.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>C. F.</given-names>
</name>
</person-group> (<year>2007</year>). <article-title>A new method to measure the semantic similarity of GO terms</article-title>. <source>Bioinformatics</source> <volume>23</volume>, <fpage>1274</fpage>&#x2013;<lpage>1281</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btm087</pub-id>
<pub-id pub-id-type="pmid">17344234</pub-id>
</mixed-citation>
</ref>
<ref id="B43">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wijesooriya</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Jadaan</surname>
<given-names>S. A.</given-names>
</name>
<name>
<surname>Perera</surname>
<given-names>K. L.</given-names>
</name>
<name>
<surname>Kaur</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Ziemann</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Urgent need for consistent standards in functional enrichment analysis</article-title>. <source>PLoS Comput. Biol.</source> <volume>18</volume>, <fpage>e1009935</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pcbi.1009935</pub-id>
<pub-id pub-id-type="pmid">35263338</pub-id>
</mixed-citation>
</ref>
<ref id="B44">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xin</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Dang</surname>
<given-names>L. T.</given-names>
</name>
<name>
<surname>Burke</surname>
<given-names>H. M. S.</given-names>
</name>
<name>
<surname>Revote</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Charitakis</surname>
<given-names>N.</given-names>
</name>
<etal/>
</person-group> (<year>2022</year>). <article-title>MonaGO: a novel gene ontology enrichment analysis visualisation system</article-title>. <source>BMC Bioinforma.</source> <volume>23</volume>, <fpage>69</fpage>. <pub-id pub-id-type="doi">10.1186/s12859-022-04594-1</pub-id>
<pub-id pub-id-type="pmid">35164667</pub-id>
</mixed-citation>
</ref>
<ref id="B45">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xu</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Luo</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Zhan</surname>
<given-names>L.</given-names>
</name>
<etal/>
</person-group> (<year>2024</year>). <article-title>Using clusterProfiler to characterize multiomics data</article-title>. <source>Nat. Protoc.</source> <volume>19</volume> (<issue>11</issue>), <fpage>3292</fpage>&#x2013;<lpage>3320</lpage>. <pub-id pub-id-type="doi">10.1038/s41596-024-01020-z</pub-id>
<pub-id pub-id-type="pmid">39019974</pub-id>
</mixed-citation>
</ref>
<ref id="B46">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ziemann</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Schroeter</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Bora</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Two subtle problems with overrepresentation analysis</article-title>. <source>Bioinforma. Adv.</source> <volume>4</volume>, <fpage>vbae159</fpage>. <pub-id pub-id-type="doi">10.1093/BIOADV/VBAE159</pub-id>
<pub-id pub-id-type="pmid">39539946</pub-id>
</mixed-citation>
</ref>
<ref id="B47">
<mixed-citation publication-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zyla</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Marczyk</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Domaszewska</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Kaufmann</surname>
<given-names>S. H. E.</given-names>
</name>
<name>
<surname>Polanska</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Weiner</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Gene set enrichment for reproducible science: comparison of CERNO and eight other algorithms</article-title>. <source>Bioinformatics</source> <volume>35</volume>, <fpage>5146</fpage>&#x2013;<lpage>5154</lpage>. <pub-id pub-id-type="doi">10.1093/bioinformatics/btz447</pub-id>
<pub-id pub-id-type="pmid">31165139</pub-id>
</mixed-citation>
</ref>
</ref-list>
</back>
</article>