<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Neuroinform.</journal-id>
<journal-title>Frontiers in Neuroinformatics</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Neuroinform.</abbrev-journal-title>
<issn pub-type="epub">1662-5196</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fninf.2023.1123376</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Neuroscience</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Novel methods for elucidating modality importance in multimodal electrophysiology classifiers</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name><surname>Ellis</surname> <given-names>Charles A.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/934990/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Sendi</surname> <given-names>Mohammad S. E.</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/930799/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Zhang</surname> <given-names>Rongen</given-names></name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2140207/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Carbajal</surname> <given-names>Darwin A.</given-names></name>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
</contrib>
<contrib contrib-type="author">
<name><surname>Wang</surname> <given-names>May D.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1247221/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Miller</surname> <given-names>Robyn L.</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="aff" rid="aff6"><sup>6</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/197665/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Calhoun</surname> <given-names>Vince D.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="aff" rid="aff6"><sup>6</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/884/overview"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>The Wallace H. Coulter Department of Biomedical Engineering, Georgia Institute of Technology, Emory University</institution>, <addr-line>Atlanta, GA</addr-line>, <country>United States</country></aff>
<aff id="aff2"><sup>2</sup><institution>Tri-Institutional Center for Translational Research in Neuroimaging and Data Science, Georgia State University, Georgia Institute of Technology, Emory University</institution>, <addr-line>Atlanta, GA</addr-line>, <country>United States</country></aff>
<aff id="aff3"><sup>3</sup><institution>McLean Hospital and Harvard Medical School</institution>, <addr-line>Boston, MA</addr-line>, <country>United States</country></aff>
<aff id="aff4"><sup>4</sup><institution>Hankamer School of Business, Baylor University</institution>, <addr-line>Waco, TX</addr-line>, <country>United States</country></aff>
<aff id="aff5"><sup>5</sup><institution>The Wallace H. Coulter Department of Biomedical Engineering, Georgia Institute of Technology</institution>, <addr-line>Atlanta, GA</addr-line>, <country>United States</country></aff>
<aff id="aff6"><sup>6</sup><institution>Department of Computer Science, Georgia State University</institution>, <addr-line>Atlanta, GA</addr-line>, <country>United States</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: Kais Gadhoumi, Duke University, United States</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Gary N. Garcia-Molina, Sleep Number Labs, United States; Etienne Thoret, Aix-Marseille Universit&#x00E9;, France</p></fn>
<corresp id="c001">&#x002A;Correspondence: Charles A. Ellis, <email>cae67@gatech.edu</email></corresp>
</author-notes>
<pub-date pub-type="epub">
<day>15</day>
<month>03</month>
<year>2023</year>
</pub-date>
<pub-date pub-type="collection">
<year>2023</year>
</pub-date>
<volume>17</volume>
<elocation-id>1123376</elocation-id>
<history>
<date date-type="received">
<day>14</day>
<month>12</month>
<year>2022</year>
</date>
<date date-type="accepted">
<day>01</day>
<month>03</month>
<year>2023</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2023 Ellis, Sendi, Zhang, Carbajal, Wang, Miller and Calhoun.</copyright-statement>
<copyright-year>2023</copyright-year>
<copyright-holder>Ellis, Sendi, Zhang, Carbajal, Wang, Miller and Calhoun</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<sec>
<title>Introduction</title>
<p>Multimodal classification is increasingly common in electrophysiology studies. Many studies use deep learning classifiers with raw time-series data, which makes explainability difficult, and has resulted in relatively few studies applying explainability methods. This is concerning because explainability is vital to the development and implementation of clinical classifiers. As such, new multimodal explainability methods are needed.</p>
</sec>
<sec>
<title>Methods</title>
<p>In this study, we train a convolutional neural network for automated sleep stage classification with electroencephalogram (EEG), electrooculogram, and electromyogram data. We then present a global explainability approach that is uniquely adapted for electrophysiology analysis and compare it to an existing approach. We present the first two local multimodal explainability approaches. We look for subject-level differences in the local explanations that are obscured by global methods and look for relationships between the explanations and clinical and demographic variables in a novel analysis.</p>
</sec>
<sec>
<title>Results</title>
<p>We find a high level of agreement between methods. We find that EEG is globally the most important modality for most sleep stages and that subject-level differences in importance arise in local explanations that are not captured in global explanations. We further show that sex, followed by medication and age, had significant effects upon the patterns learned by the classifier.</p>
</sec>
<sec>
<title>Discussion</title>
<p>Our novel methods enhance explainability for the growing field of multimodal electrophysiology classification, provide avenues for the advancement of personalized medicine, yield unique insights into the effects of demographic and clinical variables upon classifiers, and help pave the way for the implementation of multimodal electrophysiology clinical classifiers.</p>
</sec>
</abstract>
<kwd-group>
<kwd>multimodal classification</kwd>
<kwd>explainable deep learning</kwd>
<kwd>sleep stage classification</kwd>
<kwd>electrophysiology</kwd>
<kwd>electroencephalography</kwd>
<kwd>electrooculography</kwd>
<kwd>electromyography</kwd>
</kwd-group>
<contract-sponsor id="cn001">National Institutes of Health<named-content content-type="fundref-id">10.13039/100000002</named-content></contract-sponsor>
<counts>
<fig-count count="8"/>
<table-count count="1"/>
<equation-count count="9"/>
<ref-count count="85"/>
<page-count count="14"/>
<word-count count="11546"/>
</counts>
</article-meta>
</front>
<body>
<sec id="S1" sec-type="intro">
<title>1. Introduction</title>
<p>Biomedical informatics studies (<xref ref-type="bibr" rid="B43">Lin et al., 2019</xref>; <xref ref-type="bibr" rid="B45">Mellem et al., 2020</xref>; <xref ref-type="bibr" rid="B84">Zhai et al., 2020</xref>), and electrophysiology studies (<xref ref-type="bibr" rid="B52">Niroshana et al., 2019</xref>; <xref ref-type="bibr" rid="B57">Phan et al., 2019</xref>; <xref ref-type="bibr" rid="B81">Wang et al., 2020</xref>; <xref ref-type="bibr" rid="B42">Li et al., 2021</xref>) in particular, have increasingly begun to incorporate multimodal data when training machine learning classifiers. Using complementary modalities can enable the extraction of better features and improve classification performance (<xref ref-type="bibr" rid="B81">Wang et al., 2020</xref>; <xref ref-type="bibr" rid="B84">Zhai et al., 2020</xref>). While multimodal data can improve classifier performance, it can also make explaining models more difficult. This is especially true for state-of-the-art deep learning models. As a result, most studies have not used explainability (<xref ref-type="bibr" rid="B85">Zhang et al., 2011</xref>; <xref ref-type="bibr" rid="B39">Kwon et al., 2018</xref>; <xref ref-type="bibr" rid="B52">Niroshana et al., 2019</xref>; <xref ref-type="bibr" rid="B57">Phan et al., 2019</xref>; <xref ref-type="bibr" rid="B81">Wang et al., 2020</xref>; <xref ref-type="bibr" rid="B42">Li et al., 2021</xref>), which is concerning because transparency is increasingly required to assist with model development and physician decision making (<xref ref-type="bibr" rid="B71">Sullivan and Schweikart, 2019</xref>). As such, more multimodal explainability methods need to be developed (<xref ref-type="bibr" rid="B43">Lin et al., 2019</xref>; <xref ref-type="bibr" rid="B45">Mellem et al., 2020</xref>; <xref ref-type="bibr" rid="B21">Ellis et al., 2021a</xref>,<xref ref-type="bibr" rid="B29">b</xref>,<xref ref-type="bibr" rid="B22">c</xref>,<xref ref-type="bibr" rid="B26">d</xref>). In this study, we use automated sleep stage classification as a testbed for the development of multimodal explainability methods. We further present 3 novel approaches that offer significant improvements over existing approaches for use with multimodal electrophysiology data. Specifically, we present a global ablation approach that is uniquely adapted for electrophysiology data. We further present two local methods that can be used to identify personalized electrophysiology biomarkers that would be obscured by global methods. Using the local methods, we perform a novel analysis that illuminates the effects of demographic and clinical variables upon the patterns learned by the classifier.</p>
<sec id="S1.SS1">
<title>1.1. Automated sleep stage classification as testbed for multimodal explainability</title>
<p>Automated sleep stage classification offers a unique testbed for the development of novel multimodal explainability methods. Automated sleep stage classification has multiple noteworthy characteristics. (1) In practice, clinicians rely on multiple modalities instead of a single modality to manually score sleep stages (<xref ref-type="bibr" rid="B34">Iber et al., 2007</xref>). (2) The features differentiating sleep stages and the importance of modalities are well-characterized in a clinical setting (<xref ref-type="bibr" rid="B34">Iber et al., 2007</xref>). (3) Multiple large sleep stage datasets are publicly available (<xref ref-type="bibr" rid="B59">Quan et al., 1997</xref>; <xref ref-type="bibr" rid="B35">Kemp et al., 2000</xref>; <xref ref-type="bibr" rid="B36">Khalighi et al., 2016</xref>). (4) A number of studies involving unimodal and multimodal sleep stage classification have been conducted (<xref ref-type="bibr" rid="B60">Rahman et al., 2018</xref>; <xref ref-type="bibr" rid="B81">Wang et al., 2020</xref>), which could enable data scientists to develop their explainability methods alongside established architectures. Because these characteristics can help us validate our explainability methods and because there is a clinical need for explainability in sleep stage classification, we chose sleep stage classification as a use-case in this study. In the following paragraphs, we briefly review the domain of sleep stage classification and the explainability methods that have been used within the domain, both for unimodal and multimodal classification. A description of sleep stages can be found in the <xref ref-type="supplementary-material" rid="DS1">Supplementary section</xref>, &#x201C;Characteristic Features of Sleep Stages.&#x201D;</p>
</sec>
<sec id="S1.SS2">
<title>1.2. Unimodal sleep stage classification and explainability</title>
<p>Typical sleep stage classification approaches involve the classification of 5 stages: Awake, rapid eye movement (REM), non-REM1 (NREM1), NREM2, and NREM3. Many sleep stage classification studies have used unimodal EEG. Some studies have used extracted features for sleep stage classification (<xref ref-type="bibr" rid="B2">Aboalayon et al., 2015</xref>; <xref ref-type="bibr" rid="B63">Rojas et al., 2017</xref>; <xref ref-type="bibr" rid="B60">Rahman et al., 2018</xref>; <xref ref-type="bibr" rid="B46">Michielli et al., 2019</xref>), but recent studies have begun to use deep learning methods involving automated feature extraction from raw data (<xref ref-type="bibr" rid="B76">Tsinalis et al., 2016a</xref>; <xref ref-type="bibr" rid="B63">Rojas et al., 2017</xref>; <xref ref-type="bibr" rid="B72">Supratak et al., 2017</xref>; <xref ref-type="bibr" rid="B69">Sors et al., 2018</xref>; <xref ref-type="bibr" rid="B50">Mousavi et al., 2019</xref>; <xref ref-type="bibr" rid="B20">Eldele et al., 2021</xref>). Multiple recent studies have involved explainability methods. In a couple of studies, authors trained convolutional neural networks (CNNs) to classify EEG spectrograms and applied sensitivity or activation maximization (<xref ref-type="bibr" rid="B68">Simonyan et al., 2013</xref>) to identify the important features (<xref ref-type="bibr" rid="B80">Vilamala et al., 2017</xref>; <xref ref-type="bibr" rid="B64">Ruffini et al., 2019</xref>). In other studies, authors trained interpretable machine learning models or deep learning models with layer-wise relevance propagation (LRP) (<xref ref-type="bibr" rid="B8">Bach et al., 2015</xref>) to classify power spectral density values and gain insight into the features learned by the classifiers (<xref ref-type="bibr" rid="B16">Chen et al., 2019</xref>; <xref ref-type="bibr" rid="B28">Ellis et al., 2021e</xref>). A few studies involving deep learning models with raw data have also used explainability methods (<xref ref-type="bibr" rid="B50">Mousavi et al., 2019</xref>; <xref ref-type="bibr" rid="B23">Ellis et al., 2021f</xref>,<xref ref-type="bibr" rid="B24">g</xref>,<xref ref-type="bibr" rid="B27">h</xref>). These studies typically seek to identify the spectral features (<xref ref-type="bibr" rid="B51">Nahmias and Kontson, 2020</xref>; <xref ref-type="bibr" rid="B9">Barnes et al., 2021</xref>; <xref ref-type="bibr" rid="B23">Ellis et al., 2021f</xref>,<xref ref-type="bibr" rid="B24">g</xref>,<xref ref-type="bibr" rid="B27">h</xref>,<xref ref-type="bibr" rid="B25">i</xref>) or waveforms (<xref ref-type="bibr" rid="B27">Ellis et al., 2021h</xref>,<xref ref-type="bibr" rid="B25">i</xref>) learned by neural networks. However, multimodal classification poses unique challenges for explainability that do not exist for unimodal classification.</p>
</sec>
<sec id="S1.SS3">
<title>1.3. Multimodal explainability in sleep stage classification and other domains</title>
<p>Most multimodal classification studies, regardless of whether they used extracted features (<xref ref-type="bibr" rid="B57">Phan et al., 2019</xref>; <xref ref-type="bibr" rid="B42">Li et al., 2021</xref>) or raw data (<xref ref-type="bibr" rid="B52">Niroshana et al., 2019</xref>; <xref ref-type="bibr" rid="B81">Wang et al., 2020</xref>), have not used explainability methods. Among the few studies involving explainability (<xref ref-type="bibr" rid="B40">Lajnef et al., 2015</xref>; <xref ref-type="bibr" rid="B14">Chambon et al., 2018</xref>; <xref ref-type="bibr" rid="B54">Pathak et al., 2021</xref>), some have used extracted features and forward feature selection (FFS) (<xref ref-type="bibr" rid="B40">Lajnef et al., 2015</xref>). Others have used raw data and ablation for insight into modality importance (<xref ref-type="bibr" rid="B54">Pathak et al., 2021</xref>). Additionally, some have shown the importance of EEG spectra or performance increases after retraining a model with additional modalities (<xref ref-type="bibr" rid="B14">Chambon et al., 2018</xref>). Some multimodal explainability methods are also found in other domains (<xref ref-type="bibr" rid="B43">Lin et al., 2019</xref>; <xref ref-type="bibr" rid="B45">Mellem et al., 2020</xref>; <xref ref-type="bibr" rid="B58">Porumb et al., 2020</xref>). Similar to (<xref ref-type="bibr" rid="B40">Lajnef et al., 2015</xref>), one paper used FFS to find key features from clinical scales and imaging features (<xref ref-type="bibr" rid="B45">Mellem et al., 2020</xref>). One study used impurity and ablation (<xref ref-type="bibr" rid="B43">Lin et al., 2019</xref>). Another study identified important time windows in one modality (<xref ref-type="bibr" rid="B58">Porumb et al., 2020</xref>) with Grad-CAM (<xref ref-type="bibr" rid="B67">Selvaraju et al., 2020</xref>).</p>
</sec>
<sec id="S1.SS4">
<title>1.4. Existing multimodal explainability methods</title>
<p>As previously described, multiple explainability methods have been used with multimodal classifiers: FFS (<xref ref-type="bibr" rid="B40">Lajnef et al., 2015</xref>; <xref ref-type="bibr" rid="B45">Mellem et al., 2020</xref>), impurity (<xref ref-type="bibr" rid="B43">Lin et al., 2019</xref>), and ablation (<xref ref-type="bibr" rid="B43">Lin et al., 2019</xref>; <xref ref-type="bibr" rid="B54">Pathak et al., 2021</xref>). FFS is applicable to most classifiers. However, it requires retraining models many times, which is impractical for computationally intensive deep learning frameworks. Impurity is only applicable to tree-based classifiers. Lastly, ablation is, like FFS, also applicable to nearly any classifier and is easy to implement. In contrast to FFS, ablation is not computationally intensive. As such, of existing approaches, it is most useful for finding modality importance in deep learning classifiers.</p>
</sec>
<sec id="S1.SS5">
<title>1.5. Limitations of existing ablation approaches and novel alternatives</title>
<p>Ablation is related to perturbation-based methods like RISE (<xref ref-type="bibr" rid="B55">Petsiuk et al., 2018</xref>) or LIME (<xref ref-type="bibr" rid="B62">Ribeiro et al., 2016</xref>) that are frequently used in explainability for image classification and to methods like those presented in <xref ref-type="bibr" rid="B75">Thoret et al. (2021)</xref> that have been used in neuroscience applications. Importantly, ablation has a key weakness like all perturbation-based explainability methods. Specifically, perturbation methods can create out-of-distribution samples that lead to a poor estimates of modality importance (<xref ref-type="bibr" rid="B47">Molnar, 2018</xref>). Ablation involves (1) The substitution of a modality with neutral values (i.e., that do not give evidence for any one class) and (2) an examination of how that ablation affects the classifier. As such, when translating ablation to a new domain, it is important to consider how to set a modality to a neutral state while minimizing the likelihood out-of-distribution samples and features creation. Existing studies using ablation have replaced each modality with zeros (<xref ref-type="bibr" rid="B43">Lin et al., 2019</xref>; <xref ref-type="bibr" rid="B54">Pathak et al., 2021</xref>). However, zeroing out modalities creates samples that are highly irregular within the electrophysiology domain. In contrast, electrodes commonly return some line-related noise or, in instances when an electrode is not working properly, only return line-related noise. Line-related noise is found in electrophysiology data at 50 or 60 Hz due to the presence of lights, power lines and other electronics near recording devices. Because it is so often found in electrophysiology data, a classifier should learn to ignore it, and it should be neutral to the classifier. As such, line-related noise could offer a more reliable, electrophysiology-specific alternative to the zero-out ablation methods that have previously been applied.</p>
<p>While line noise-based ablation would be less likely to produce out-of-distribution samples and features than a zero-out ablation approach, it would still be at risk of doing so. Gradient-based feature attribution (GBFA) methods (<xref ref-type="bibr" rid="B4">Ancona et al., 2018</xref>) like Grad-CAM (<xref ref-type="bibr" rid="B67">Selvaraju et al., 2020</xref>), saliency (<xref ref-type="bibr" rid="B68">Simonyan et al., 2013</xref>), and LRP (<xref ref-type="bibr" rid="B8">Bach et al., 2015</xref>), in particular, offer an alternative to ablation that does not risk producing out-of-distribution samples. Additionally, local ablation methods, similar to saliency (<xref ref-type="bibr" rid="B68">Simonyan et al., 2013</xref>), show what features or time points make a sample more or less like the patterns learned by the classifier for a particular class. LRP shows what features or time points are actually used by the classifier for its classification and indicates their importance (<xref ref-type="bibr" rid="B48">Montavon et al., 2018</xref>).</p>
</sec>
<sec id="S1.SS6">
<title>1.6. Limitations of global explanations and proposal of novel local explainability approach</title>
<p>Global explainability methods identify the general importance of each modality to the classifier. In contrast, local methods provide higher resolution insight and indicate the importance of each modality to the classification of individual samples (<xref ref-type="bibr" rid="B47">Molnar, 2018</xref>). Global methods have inherent limitations relative to local methods, and existing multimodal explainability approaches have mainly been global. Importantly, as shown in <xref ref-type="fig" rid="F1">Figure 1</xref>, global explanations obscure feature importance for individual samples and can obscure the presence of subgroups. Local explanations for many samples can be combined for higher level or global importance estimates (<xref ref-type="bibr" rid="B28">Ellis et al., 2021e</xref>,<xref ref-type="bibr" rid="B23">f</xref>). Because of this, they can also be analyzed on a subject-specific level that paves the way for the identification of personalized biomarkers. Furthermore, local explanations can be used to examine the degree to which demographic and clinical variables affect the patterns learned by a classifier for specific classes and features (<xref ref-type="bibr" rid="B22">Ellis et al., 2021c</xref>), which is a capacity that has not previously been exploited in multimodal classification. Local methods have been applied in a couple multimodal classification studies. In one study, authors ablated time points of an input sample and examined the effect on the classification of the sample (<xref ref-type="bibr" rid="B54">Pathak et al., 2021</xref>). In another study, authors used Grad-CAM to examine segments of a single modality (<xref ref-type="bibr" rid="B58">Porumb et al., 2020</xref>). Neither study identified the importance of each modality.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption><p>Example of global versus local importance with dummy data. From left to right, are 4 importance metrics: the importance for an atypical sample, the importance for two generic samples, and the global importance estimate that is formed by averaging the importance values for the individual samples. Note that modality two is very important for the atypical subject but not for the generic subject, so the presence of the atypical sample is hidden in the global importance. When thousands of samples from dozens of subjects are being analyzed, the presence of subgroups is easily obscured.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fninf-17-1123376-g001.tif"/>
</fig>
<p>In the present study, we train a CNN for automated sleep stage classification using a publicly available dataset. We introduce a global ablation approach that is uniquely adapted for the electrophysiology domain (<xref ref-type="bibr" rid="B29">Ellis et al., 2021b</xref>). We then present a local ablation approach (<xref ref-type="bibr" rid="B22">Ellis et al., 2021c</xref>) and show how GBFA methods can be used for local insight into multimodal classifiers (<xref ref-type="bibr" rid="B21">Ellis et al., 2021a</xref>). With our local methods, we identify subject-level differences in modality importance that support the viability of the methods for personalized biomarker identification. We then use the local explanations in a novel analysis that provides insight into the patterns learned by the classifier related to the age, sex, and state of medication of subjects in our dataset (<xref ref-type="bibr" rid="B22">Ellis et al., 2021c</xref>,<xref ref-type="bibr" rid="B26">d</xref>).</p>
</sec>
</sec>
<sec id="S2" sec-type="materials|methods">
<title>2. Materials and methods</title>
<p>In this section, we describe our data, preprocessing, model architecture and training approach, and explainability methods.</p>
<sec id="S2.SS1">
<title>2.1. Description of data</title>
<p>We utilized Sleep Telemetry data from the Sleep-EDF Expanded Database (<xref ref-type="bibr" rid="B35">Kemp et al., 2000</xref>) on Physionet (<xref ref-type="bibr" rid="B32">Goldberger et al., 2000</xref>). The database has been used in previous studies (<xref ref-type="bibr" rid="B80">Vilamala et al., 2017</xref>; <xref ref-type="bibr" rid="B60">Rahman et al., 2018</xref>; <xref ref-type="bibr" rid="B50">Mousavi et al., 2019</xref>; <xref ref-type="bibr" rid="B57">Phan et al., 2019</xref>). Because the dataset was publicly available, no Internal Review Board approval was needed. The dataset has 44 approximately 9-h recordings from 22 subjects (15 female and 7 male) with primary sleep onset insomnia (<xref ref-type="bibr" rid="B78">Tuk et al., 1997</xref>). Subject age had a mean of 40.18 years and a standard deviation of 18.09 years. <xref ref-type="fig" rid="F2">Figure 2</xref> shows subject demographics. All subjects had two recordings&#x2013;one following placebo administration and one following temazepam administration. Temazepam belongs to a class of drugs called benzodiazepines which amplify the effects of the neurotransmitter y-aminobutyric acid (GABA). GABA is inhibitory in nature and produces a calming effect on the brain (<xref ref-type="bibr" rid="B33">Griffin et al., 2013</xref>). It is often used to treat insomnia and affects electrophysiology activity. Each recording had data from 4 electrodes: 2 EEG, 1 EOG, and 1 EMG. Data was recorded at a 100 Hertz (Hz) sampling frequency. The EEG electrodes were FPz-Cz and Pz-Oz (<xref ref-type="bibr" rid="B79">Van Sweden et al., 1990</xref>), but like previous studies (<xref ref-type="bibr" rid="B77">Tsinalis et al., 2016b</xref>; <xref ref-type="bibr" rid="B80">Vilamala et al., 2017</xref>; <xref ref-type="bibr" rid="B46">Michielli et al., 2019</xref>; <xref ref-type="bibr" rid="B50">Mousavi et al., 2019</xref>; <xref ref-type="bibr" rid="B57">Phan et al., 2019</xref>), we used only Fpz-Cz. A 1-Hz marker indicated the presence of recording errors. Using the Rechtschaffen and Kales standard (<xref ref-type="bibr" rid="B61">Rechtschaffen and Kales, 1968</xref>), experts assigned 30-s epochs to seven categories: Movement, Awake, REM, NREM1, NREM2, NREM3, and NREM4. We merged NREM3 and NREM4 into a single NREM3 class (<xref ref-type="bibr" rid="B34">Iber et al., 2007</xref>), and we removed all samples containing movement or recording errors.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption><p>Distribution of samples and subject demographics. Panels <bold>(A,B)</bold> show the distributions of temazepam and placebo samples, respectively, for each subject. Panel <bold>(C)</bold> shows the age and sex of each subject, with the subjects arranged from youngest to oldest. Each panel shares the same <italic>x</italic>-axis.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fninf-17-1123376-g002.tif"/>
</fig>
</sec>
<sec id="S2.SS2">
<title>2.2. Description of data preprocessing</title>
<p>Based on the data annotation, we segmented the data into 30-s samples. Within each recording, we separately z-scored each electrode to improve cross-subject pattern identification. Our final dataset had 42,218 samples. The dataset was highly imbalanced with Awake, NREM1, NREM2, NREM3, and REM classes having 9.97, 8.53, 46.8, 14.92, and 19.78% of the dataset, respectively. We did not perform any filtering or reject data due to quality or noise issues.</p>
</sec>
<sec id="S2.SS3">
<title>2.3. Description of 1D-CNN</title>
<sec id="S2.SS3.SSS1">
<title>2.3.1. Model architecture and training</title>
<p>We adapted a CNN architecture initially developed for EEG classification (<xref ref-type="bibr" rid="B83">Youness, 2020</xref>). The architecture is shown in <xref ref-type="fig" rid="F3">Figure 3</xref>. We implemented the architecture in Keras (<xref ref-type="bibr" rid="B18">Chollet, 2015</xref>) with a TensorFlow (<xref ref-type="bibr" rid="B1">Abadi et al., 2016</xref>) backend. We used 10-fold cross-validation with a random 17-2-3 subject training-validation-test split each fold. We used class-weighted categorical cross entropy loss to account for class imbalances. We used a batch size of 100 with shuffling after each epoch. We used the Adam optimizer (<xref ref-type="bibr" rid="B38">Kingma and Ba, 2015</xref>) with an adaptive learning rate. Starting at a learning rate of 0.001, the step size decreased by a factor of 10 if validation accuracy did not improve within a 5-epoch window. We used early stopping to end training if validation accuracy plateaued for 20 epochs with a maximum of 100 epochs and used model checkpoints to select the model from each fold that obtained the best validation accuracy. We used the selected models for testing and explainability.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption><p>CNN architecture. Layers (i) of the diagram repeat 3 times. In (i) there are 6 1D-convolutional (conv1d) layers in total. The first two conv1d layers (number of filters = 16, kernel size = 5) are followed by max pooling (pool size = 2) and spatial dropout (rate = 0.01). The second two conv1d layers (number of filters = 32, kernel size = 3) are followed by max pooling (pool size = 2) and a spatial dropout (rate = 0.01). The third pair of conv1d layers (number of filters = 32, kernel size = 3) are followed by max pooling (pool size = 2) and spatial dropout (rate = 0.01). In (ii), the last two conv1d layers (number of filters = 256, kernel size = 3) are followed by global max pooling, dropout (rate = 0.01), and a flatten layer. The first two dense layers (number of nodes = 64) have dropout rates of 0.1 and 0.05, respectively. The last dense layer has 5 nodes. The inputs sample shape was 3,000 time points &#x00D7; 3 modalities, and no layers used zero padding during convolution. An &#x201C;R&#x201D; or an &#x201C;S&#x201D; indicates that a layer is followed by ReLU or Softmax activation functions, respectively.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fninf-17-1123376-g003.tif"/>
</fig>
</sec>
<sec id="S2.SS3.SSS2">
<title>2.3.2. Model performance evaluation</title>
<p>When evaluating model test performance, we sought to account for class imbalances. We calculated the precision, recall, and F1 score for each class. We calculated the mean and standard deviation of the metrics across folds.</p>
<disp-formula id="S2.Ex1">
<mml:math id="M1">
<mml:mrow>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>c</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mpadded width="+3.3pt">
<mml:mi>n</mml:mi>
</mml:mpadded>
</mml:mrow>
<mml:mo rspace="5.8pt">=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>u</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>v</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>u</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>v</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mo>+</mml:mo>
<mml:mrow>
<mml:mi>F</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>v</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="S2.Ex2">
<mml:math id="M2">
<mml:mrow>
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>c</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mpadded width="+3.3pt">
<mml:mi>l</mml:mi>
</mml:mpadded>
</mml:mrow>
<mml:mo rspace="5.8pt">=</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>u</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>v</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mi>T</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>u</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>v</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
</mml:mrow>
<mml:mo>+</mml:mo>
<mml:mrow>
<mml:mi>F</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>N</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>g</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>v</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="S2.Ex3">
<mml:math id="M3">
<mml:mrow>
<mml:mrow>
<mml:mi>F</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mpadded width="+3.3pt">
<mml:mn>1</mml:mn>
</mml:mpadded>
</mml:mrow>
<mml:mo rspace="5.8pt">=</mml:mo>
<mml:mrow>
<mml:mn>2</mml:mn>
<mml:mo>&#x002A;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>c</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>n</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mmultiscripts>
<mml:mi>R</mml:mi>
<mml:mprescripts/>
<mml:none/>
<mml:mo>&#x002A;</mml:mo>
</mml:mmultiscripts>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>c</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>c</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>n</mml:mi>
</mml:mrow>
<mml:mo>+</mml:mo>
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>c</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mrow>
</mml:math>
</disp-formula>
</sec>
</sec>
<sec id="S2.SS4">
<title>2.4. Description of global ablation approaches</title>
<p>We applied two global ablation approaches to estimate class-specific modality importance. We presented a novel global ablation approach that is uniquely adapted to the electrophysiology domain (<xref ref-type="bibr" rid="B29">Ellis et al., 2021b</xref>) and compared our approach to a standard approach that has been used in previous studies (<xref ref-type="bibr" rid="B43">Lin et al., 2019</xref>; <xref ref-type="bibr" rid="B54">Pathak et al., 2021</xref>).</p>
<p>Generally, ablation takes place after model training. It involves replacing a feature or modality with zeros during model evaluation and examining the change in model performance following the loss of the information in that modality or feature. The importance of the replaced feature or modality to the model is directly related to the decrease in model performance associated with its loss. A feature <italic>f1</italic> is more important to a model than a feature <italic>f2</italic> if the effect of ablating <italic>f1</italic> is greater than the effect of ablating <italic>f2.</italic> Our standard ablation approach had several key steps (1). We calculated a confusion matrix (i.e., a model performance estimate) for the test data in a fold (2). We replaced a modality <italic>m</italic> with zeros across all test samples in the fold (i.e., ablation) (3). We calculated a confusion matrix for the classifier on the test data with the replaced modality <italic>m</italic> (4). We calculated the percent change (PCG) in samples assigned to each classification group following ablation (i.e., effect of ablation). Example classification groups include NREM1 samples classified as REM, NREM2 samples classified as NREM3, and REM samples classified as REM (5). We repeated steps 2 through 4 for each modality <italic>m</italic> (6). We repeated steps 1 through 5 for each fold.</p>
<disp-formula id="S2.Ex4">
<mml:math id="M4">
<mml:mrow>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>C</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mpadded width="+3.3pt">
<mml:mi>G</mml:mi>
</mml:mpadded>
</mml:mrow>
<mml:mo>=</mml:mo>
<mml:mi/>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="S2.Ex5">
<mml:math id="M6">
<mml:mrow>
<mml:mn>100</mml:mn>
<mml:mo>&#x002A;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>u</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>M</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>p</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
</mml:mrow>
<mml:mo>-</mml:mo>
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>u</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>U</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>n</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>p</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>u</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>U</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>n</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>p</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mstyle>
</mml:mrow>
</mml:math>
</disp-formula>
<p>We propose an ablation approach for multimodal electrophysiology analysis that involves replacing modalities in a way that mimics line-related noise. This approach involves all of the steps detailed previously. However, we modify Step 2 of the ablation process. Instead of replacing modality <italic>m</italic> with zeros, we replace modality <italic>m</italic> with a combination of a sinusoid and Gaussian noise. We use a sinusoid with a frequency of 50 Hz and an amplitude of 0.1, and the Gaussian noise had a mean of 0 and standard deviation of 0.1. To determine whether our line-related noise approach yielded results significantly different from standard ablation, we performed a series of two-tailed t-tests. Within each modality, we compared the importance values in each classification group for each method across folds.</p>
</sec>
<sec id="S2.SS5">
<title>2.5. Description of novel local ablation approach</title>
<p>We developed a local ablation approach for insight into modality importance (<xref ref-type="bibr" rid="B22">Ellis et al., 2021c</xref>). Our novel ablation approach is similar to the global approach described in the previous section (1). We obtained the top-class probability for a sample (2). We ablated a modality in that sample (3). We obtained the classification probability of the modified sample for the original top class (4). We computed the percent change in classification probability (5). We repeated steps 2 through 4 for each modality (6). We repeated steps 2 through 5 for each sample (7). We repeated steps 2 through 6 for each fold.</p>
<disp-formula id="S2.Ex6">
<mml:math id="M8">
<mml:mrow>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>C</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mpadded width="+3.3pt">
<mml:mi>G</mml:mi>
</mml:mpadded>
</mml:mrow>
<mml:mo>=</mml:mo>
<mml:mi/>
</mml:mrow>
</mml:math>
</disp-formula>
<disp-formula id="S2.Ex7">
<mml:math id="M10">
<mml:mrow>
<mml:mn>100</mml:mn>
<mml:mo>&#x002A;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>p</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mpadded width="+5pt">
<mml:mi>y</mml:mi>
</mml:mpadded>
</mml:mrow>
<mml:mo>-</mml:mo>
<mml:mrow>
<mml:mi>U</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>n</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>p</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>y</mml:mi>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mi>U</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>n</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>f</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>d</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>m</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>p</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>e</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>P</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>r</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>o</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>a</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>b</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>l</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>y</mml:mi>
</mml:mrow>
</mml:mfrac>
</mml:mstyle>
</mml:mrow>
</mml:math>
</disp-formula>
<p>Because there were no preexisting local approaches, we compared our local ablation results to the global ablation results. In addition to generating local visualizations of our results, we estimated global importance by calculating the mean absolute percent change in classification probability for each fold.</p>
</sec>
<sec id="S2.SS6">
<title>2.6. Description of layer-wise relevance propagation analysis</title>
<p>Layer-wise relevance propagation (LRP) (<xref ref-type="bibr" rid="B8">Bach et al., 2015</xref>) was first developed for image analysis but has since been used in electrophysiology (<xref ref-type="bibr" rid="B70">Sturm et al., 2016</xref>) and other neuroscience domains (<xref ref-type="bibr" rid="B82">Yan et al., 2017</xref>; <xref ref-type="bibr" rid="B73">Thomas et al., 2018</xref>; <xref ref-type="bibr" rid="B28">Ellis et al., 2021e</xref>). We implemented LRP with the Innvestigate library (<xref ref-type="bibr" rid="B3">Alber et al., 2019</xref>). LRP is a local explainability method but has can be used for global importance estimates (<xref ref-type="bibr" rid="B21">Ellis et al., 2021a</xref>,<xref ref-type="bibr" rid="B28">e</xref>). LRP involves several steps (1). A sample is passed through a network and assigned a class (2). A total relevance of 1 is placed at the output node of the assigned class (3). The relevance is propagated through the network to the input sample space with relevance rules. Importantly, the total relevance is conserved when propagated through the network such that the total relevance assigned to the sample space should equal the original total relevance. LRP can output both negative and positive relevance. Negative relevance indicates features that support a sample being classified as a class other than that which it was assigned. Positive relevance indicates features that support a sample being classified as its assigned class. In our study, we used the &#x03B5;-rule and &#x03B1;&#x03B2;-rule. The equation below shows the &#x03B5;-rule.</p>
<disp-formula id="S2.Ex8">
<mml:math id="M12">
<mml:mrow>
<mml:mpadded width="+3.3pt">
<mml:msub>
<mml:mi>R</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:mpadded>
<mml:mo rspace="5.8pt">=</mml:mo>
<mml:mrow>
<mml:munder>
<mml:mo largeop="true" movablelimits="false" symmetric="true">&#x2211;</mml:mo>
<mml:mi>k</mml:mi>
</mml:munder>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi>a</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
<mml:mo>&#x2062;</mml:mo>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="normal">&#x03B5;</mml:mi>
<mml:mo>+</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mo largeop="true" symmetric="true">&#x2211;</mml:mo>
<mml:mrow>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:msub>
<mml:mi>a</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
<mml:mo>&#x2062;</mml:mo>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#x2062;</mml:mo>
<mml:msub>
<mml:mi>R</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
</disp-formula>
<p>where <italic>k</italic> indicates a node that is one of <italic>k</italic>nodes in a layer deeper in a network and <italic>j</italic> indicates a node in the layer to which relevance is being propagated. <italic>R<sub>k</sub></italic> indicates the total relevance assigned to a node in a deeper layer, and <italic>R<sub>j</sub></italic> indicates the total relevance that will be assigned to a node in a shallower layer. The variables <italic>a<sub>j</sub></italic> and <italic>w</italic><sub><italic>jk</italic></sub> indicate the activation output of the layer j and the value of the weight connecting the node in layer <italic>j</italic> and node in layer <italic>k</italic>. The numerator indicates a portion of the effect that the node in layer <italic>j</italic> has upon the node in layer <italic>k</italic>, and the denominator indicates the total effect of all nodes in layer <italic>j</italic> upon the node in layer <italic>k</italic>. This combined with the summation &#x03A3;<sub><italic>k</italic></sub> indicates that the relevance assigned to the node in layer <italic>j</italic> is the sum of the fraction of the effect of the node in layer <italic>j</italic> upon all of the nodes in layer <italic>k</italic> multiplied by their respective relevance. The term &#x201C;&#x03B5;&#x201D; enables relevance to be filtered when propagated through the network. A larger &#x03B5; shrinks the amount of relevance propagated backward for nodes that would otherwise be assigned low relevance. In effect, this reduces the noisiness of the explanations. We used the &#x03B5;-rule with an &#x03B5; of 0.01 and 100.</p>
<p>The &#x03B1;&#x03B2;-rule is shown in the equation below,</p>
<disp-formula id="S2.Ex9">
<mml:math id="M13">
<mml:mrow>
<mml:mpadded width="+3.3pt">
<mml:msub>
<mml:mi>R</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:mpadded>
<mml:mo rspace="5.8pt">=</mml:mo>
<mml:mrow>
<mml:munder>
<mml:mo largeop="true" movablelimits="false" symmetric="true">&#x2211;</mml:mo>
<mml:mi>k</mml:mi>
</mml:munder>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mi mathvariant="normal">&#x03B1;</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mfrac>
<mml:msup>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>a</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
<mml:mo>&#x2062;</mml:mo>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>+</mml:mo>
</mml:msup>
<mml:mrow>
<mml:msub>
<mml:mo largeop="true" symmetric="true">&#x2211;</mml:mo>
<mml:mrow>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:msub>
<mml:msup>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>a</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
<mml:mo>&#x2062;</mml:mo>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>+</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
<mml:mo>-</mml:mo>
<mml:mrow>
<mml:mi mathvariant="normal">&#x03B2;</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mfrac>
<mml:msup>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>a</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
<mml:mo>&#x2062;</mml:mo>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>-</mml:mo>
</mml:msup>
<mml:mrow>
<mml:msub>
<mml:mo largeop="true" symmetric="true">&#x2211;</mml:mo>
<mml:mrow>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:msub>
<mml:msup>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>a</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
<mml:mo>&#x2062;</mml:mo>
<mml:msub>
<mml:mi>w</mml:mi>
<mml:mrow>
<mml:mi>j</mml:mi>
<mml:mo>&#x2062;</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>-</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2062;</mml:mo>
<mml:msub>
<mml:mi>R</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mrow>
</mml:mrow>
</mml:math>
</disp-formula>
<p>where the relevance is split into positive and negative portions. The variables &#x03B1; and &#x03B2; control how much positive and negative relevance are propagated backward, respectively. In our study, we only propagated positive relevance (i.e., &#x03B1; = 1, &#x03B2; = 0).</p>
<p>In our analysis, we generated a global estimation of importance by calculating the percent of absolute relevance assigned to each modality. We computed this value for each classification group in each fold. We also visualized how the percent of relevance varied over time.</p>
</sec>
<sec id="S2.SS7">
<title>2.7. Description of statistical analyses</title>
<p>We performed a series of statistical analyses with the local ablation and LRP (&#x03B5;-rule with &#x03B5; = 100) explanations for insight into the effects of demographic and clinical variables upon the classifier. To account for interaction effects, we trained an ordinary least squares regression model with age, medication, and sex as the independent variables and with the absolute importance (i.e., percent change in activation for local ablation and relevance for LRP) for a modality and classification group as the dependent variable. For LRP, we used the percent of absolute relevance assigned to each modality for each sample. After training the model, we obtained the resulting coefficients and <italic>p</italic>-values for each class. The sign of the coefficients identified the direction of the importance difference. After obtaining <italic>p</italic>-values, we performed false discovery rate (FDR) correction (&#x03B1; = 0.05) with the 25 <italic>p</italic>-values (i.e., 5 classes &#x00D7; 5 classes) associated with each clinical or demographic variable to account for multiple comparisons.</p>
</sec>
</sec>
<sec id="S3" sec-type="results">
<title>3. Results</title>
<p>Here, we describe our model performance, explainability, and statistical analysis results.</p>
<sec id="S3.SS1">
<title>3.1. Model performance results</title>
<p><xref ref-type="table" rid="T1">Table 1</xref> shows the mean and standard deviation of the precision, recall, and F1 score for each class. The model had highest F1 scores for NREM2 and Awake. Possibly because of its smaller sample size, NREM1 had the lowest classification performance across all metrics. While performance for NREM3 and REM was not as high as for NREM2 and Awake for most metrics, the classifier still performed well for both classes.</p>
<table-wrap position="float" id="T1">
<label>TABLE 1</label>
<caption><p>Classification performance results.</p></caption>
<table cellspacing="5" cellpadding="5" frame="box" rules="all">
<thead>
<tr>
<td valign="top" align="left" style="color:#ffffff;background-color: #7f8080;"></td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">Awake</td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">NREM1</td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">NREM2</td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">NREM3</td>
<td valign="top" align="center" style="color:#ffffff;background-color: #7f8080;">REM</td>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">F1</td>
<td valign="top" align="center">71.25 &#x00B1; 05.15</td>
<td valign="top" align="center">39.86 &#x00B1; 07.19</td>
<td valign="top" align="center">73.28 &#x00B1; 04.76</td>
<td valign="top" align="center">64.15 &#x00B1; 15.25</td>
<td valign="top" align="center">65.92 &#x00B1; 06.28</td>
</tr>
<tr>
<td valign="top" align="left">Precision</td>
<td valign="top" align="center">72.25 &#x00B1; 07.12</td>
<td valign="top" align="center">36.20 &#x00B1; 03.98</td>
<td valign="top" align="center">79.35 &#x00B1; 03.92</td>
<td valign="top" align="center">56.78 &#x00B1; 18.35</td>
<td valign="top" align="center">69.04 &#x00B1; 07.14</td>
</tr>
<tr>
<td valign="top" align="left">Recall</td>
<td valign="top" align="center">70.90 &#x00B1; 07.02</td>
<td valign="top" align="center">46.28 &#x00B1; 13.52</td>
<td valign="top" align="center">68.71 &#x00B1; 08.51</td>
<td valign="top" align="center">78.22 &#x00B1; 10.24</td>
<td valign="top" align="center">63.26 &#x00B1; 06.69</td>
</tr>
</tbody>
</table></table-wrap>
</sec>
<sec id="S3.SS2">
<title>3.2. Global explainability results</title>
<p><xref ref-type="fig" rid="F4">Figure 4</xref> and <xref ref-type="supplementary-material" rid="DS1">Supplementary Figure 1</xref> show the results comparing noise-related global ablation with the typical zero-out global ablation approach for correct classification groups and all classification groups, respectively. Interestingly, the methods generally agreed upon the relative importance of each modality. However, there were multiple significant differences in which our method seemed to amplify the effect of perturbation more than the typical zero-out approach. <xref ref-type="fig" rid="F5">Figure 5</xref> and <xref ref-type="supplementary-material" rid="DS1">Supplementary Figure 2</xref> show our local ablation results estimating global modality importance for correct classification groups and all classification groups, respectively. <xref ref-type="fig" rid="F6">Figure 6</xref> and <xref ref-type="supplementary-material" rid="DS1">Supplementary Figure 3</xref> show the LRP results for correct classification groups and all groups, respectively. We compared the relative magnitude of the estimates across methods. Across methods, EEG was generally most important. For Awake/Awake, all three methods found that EEG was most important, though LRP magnified the importance of EOG and EMG relative to EEG more than local ablation. For NREM1/NREM1, two LRP rules and local and global ablation found that EOG was most important, followed by EEG. NREM2/NREM2 and NREM3/NREM3 results were similar across methods. EEG was most important, followed by EOG and EMG. For REM/REM, only LRP &#x03B5;-rule (&#x03B5; = 0.1) agreed with local and global ablation regarding the relative modality importance. They identified the order of importance as EEG, EOG, and EMG. Many incorrect classification groups had similar distributions of relative importance across methods. However, some groups had different importance distributions. NREM1/Awake generally had greater EOG than EEG relevance for LRP but not for ablation. NREM2/NREM1 had less EOG than EEG relevance for LRP but not for ablation. Awake/REM and NREM1/REM had more EEG than EOG importance for local ablation but not for LRP. Global ablation found that EEG and EOG importance for Awake/REM varied according to the global ablation approach. Additionally, global ablation found that EEG had greater importance than EOG for NREM1/REM.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption><p>Global ablation results for correctly classified samples. The leftmost and rightmost plots show the results for zero-out and noise-related ablation, respectively. Horizontal dashed lines separate importance for each sleep stage. The <italic>x</italic>-axis indicates the mean percent change in samples across folds. Blue and red bars indicate a negative and positive percent change, respectively. Bolded modalities have significant differences (<italic>p</italic> &#x003C; 0.05) between results for the two methods. Across methods, EEG was most important for all classes except NREM1.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fninf-17-1123376-g004.tif"/>
</fig>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption><p>Local ablation results showing estimation of global importance for correct classification groups. Within each fold, we calculated the mean absolute percent change in activation for the perturbation of samples in each classification group. We then calculated the median value across folds. The <italic>x</italic>-axis indicates the values associated with that percent change in activation. The bar plots show the results for the correct classification group. EEG, EOG, and EMG importance values are shown in red, green, and blue, respectively. EEG was most important for all classes except NREM 1.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fninf-17-1123376-g005.tif"/>
</fig>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption><p>LRP global modality importance results for correct classification groups. We calculated the percent of absolute relevance for each modality across samples within each classification group. We then calculated the median value across folds. Panels <bold>(A&#x2013;C)</bold> show results for the &#x03B1;&#x03B2;-, &#x03B5;- (&#x03B5; = 100), and &#x03B5;-rules (&#x03B5; = 0.01), respectively. Red, green, and blue bars are for EEG, EOG, and EMG, respectively. EEG was generally most important.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fninf-17-1123376-g006.tif"/>
</fig>
</sec>
<sec id="S3.SS3">
<title>3.3. Subject-level local explainability results over time</title>
<p><xref ref-type="fig" rid="F7">Figure 7</xref> shows both the local ablation and LRP results over the first 2 h of a recording from Subject 12. Both methods showed similar trends in modality importance over time. They both showed lower levels of EEG and higher levels of EOG importance during Awake and NREM1 periods and showed a transition to higher EEG and lower EOG importance for NREM2, NREM3, and REM. However, LRP often seemed to more closely correspond with changes in electrophysiology activity. For example, between 60 and 80 min, EMG activity spiked, and a misclassification resulted. In this case, LRP more clearly indicated that the change affected the classification. Additionally, for NREM periods from 30 to 100 min, EEG relevance had greater variation relative to the that of other modalities than the local ablation results.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption><p>Local explanations over a 2-H sleep cycle from subject 12. The first, second, third, and fourth panels show the actual and predicted classes, electrophysiology activity, local ablation results, LRP results (&#x03B5; = 100). Unlike the global results, EOG is more important for Awake than EEG from 0 to 20 min.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fninf-17-1123376-g007.tif"/>
</fig>
</sec>
<sec id="S3.SS4">
<title>3.4. Statistical analysis of effects of clinical and demographic variables upon local explanations</title>
<p><xref ref-type="fig" rid="F8">Figures 8A&#x2013;C</xref> and <xref ref-type="supplementary-material" rid="DS1">Supplementary Figure 4</xref> show the results for the statistical analysis examining the effects of medication, sex, and age upon the local ablation explanations for correct classification groups and all groups, respectively. <xref ref-type="fig" rid="F8">Figures 8D&#x2013;F</xref> and <xref ref-type="supplementary-material" rid="DS1">Supplementary Figure 5</xref> show the results for the analysis applied to the LRP relevance. Many effects were consistent between the two methods. Across both methods, subject sex had relationships with more correct classification group modality pairs than either medication or age. The importance of EEG for Awake/Awake was less in temazepam than placebo samples and was more in temazepam than placebo samples for REM/REM. Samples assigned to NREM3 also generally had more EEG importance for temazepam than placebo samples. For EOG, most groups with significant relationships with medication had more importance across both methods in placebo than in temazepam samples. In REM/REM, EOG importance was higher in placebo than temazepam samples. NREM2/NREM1 and NREM3/NREM2 had more EMG importance in temazepam than placebo samples for both methods. REM/NREM2 had less EMG importance in temazepam than placebo samples for both methods.</p>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption><p>Effects of clinical and demographic variables for correct classification groups. Panels <bold>(A&#x2013;C)</bold> show effects of medication, sex, and age, respectively, on the local ablation results. Panels <bold>(D&#x2013;F)</bold> show effects of medication, sex, and age, respectively, on the LRP results. The <italic>x</italic>- and <italic>y</italic>-axes indicate the predicted class and modality, respectively. The heatmaps show the regression coefficient values. Bolded boxes show significant effects (<italic>p</italic> &#x003C; 0.05). A positive medication coefficient indicates that temazepam samples had more importance than placebo samples. A positive subject sex coefficient indicates that female samples had more importance than male samples, and a positive age coefficient indicates that importance increased with age. Note that sex had more significant relationships than the other variables.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fninf-17-1123376-g008.tif"/>
</fig>
<p>For the effects of sex on EEG importance, the two methods provided similar results in a couple cases: (1) NREM1/NREM2 and (2) NREM2/NREM1. However, different effects occurred in many cases: (1) NREM2/NREM2, NREM2/NREM3, and NREM2/REM, (2) NREM3/NREM3 and NREM3/NREM2, and (3) REM/REM and REM/NREM2. For EOG, correctly classified Awake, NREM2, NREM3, and REM had similar differences in importance between males and females, and the changes in importance for many classification groups were similar across methods. For EMG and sex, the differences in importance between male and female were similar in a few cases (e.g., Awake/REM, NREM1/NREM3, and REM/NREM3).</p>
<p>Across methods, age affected the importance assigned to modalities. For EEG, there were similar effects of age: (1) Awake/NREM3, (2) NREM2/NREM1 and NREM2/REM, and (3) REM/REM and REM/Awake. Many groups with different results across methods were insignificant for local ablation. For EOG, there were differences between many classification groups. However, NREM1/REM and NREM2/REM, and NREM3/NREM1 had similar results across methods. For EMG, there were many similarities in the effect of age upon the explanations of the methods: (1) NREM1/NREM2, (2) NREM2/NREM2 and NREM2/Awake, (3) NREM3/NREM3, and (4) REM/NREM3.</p>
</sec>
</sec>
<sec id="S4" sec-type="discussion">
<title>4. Discussion</title>
<p>In this section, we discuss the broader implications of our methods. We then discuss our results within the context of sleep literature and discuss future research directions.</p>
<sec id="S4.SS1">
<title>4.1. Implications of novel explainability methods beyond sleep stage classification</title>
<p>In this study, we present a series of novel multimodal explainability methods. Our global ablation method is uniquely adapted to multimodal electrophysiology data. Additionally, its finding of enhanced effects relative to a zero-out approach highlights the utility of using domain-specific perturbations. Our local ablation approach is the first local multimodal explainability method that provides insight into the importance of each modality. By examining the change in output activation following ablation, it also shows how ablation or perturbation could be used to obtain local explanations across a variety of explainability problems beyond multimodal explainability. It is, in its current state, only applicable to electrophysiology data, but it could be easily adapted to other domains. We also show, for the first time, how gradient-based methods can be used to find modality importance both locally and globally. Because they do not perturb data, GBFA methods could offer a more reliable approach than ablation. Importantly, our ablation and gradient-based methods could each be better suited to different models. Unlike gradient methods, ablation is easily applicable to all deep learning classification frameworks. For example, our ablation methods would be more effective for long short-term memory (LSTM) networks than our LRP approach. While LRP can be applied to long short-term memory networks (<xref ref-type="bibr" rid="B6">Arras et al., 2017</xref>), doing so can be challenging, especially given that LSTMs are not supported in the Innvestigate library, and problems can arise with model gradients. As such, it is generally easier to implement for most CNN or multilayer perceptron architectures (<xref ref-type="bibr" rid="B28">Ellis et al., 2021e</xref>). It is also important to note that the insights provided by ablation and LRP are slightly different. Ablation gives a quantitative estimate of the sensitivity of the model to the loss of information in a modality. In contrast, LRP gives an estimate of the reliance of the model upon a given modality in the classification of a specific sample. Our local methods could help identify subject-specific electrophysiology biomarkers for personalized medicine. Additionally, our analysis of the relationship between local explanations and demographic and clinical variables offers a way to gain insight into the effects of variables that are not explicitly included in the training data and has implications beyond multimodal explainability. Model developers could use it to better understand how aspects of data are affecting their models. Additionally, it could increase physicians&#x2019; and other relevant decision-makers&#x2019; trust of deep learning-based systems and jointly increase the likelihood of clinical adoption. It could also help scientists develop hypotheses for novel biomarkers.</p>
</sec>
<sec id="S4.SS2">
<title>4.2. Classification performance</title>
<p>Our classifier performed well but slightly below state-of-the-art classifiers (<xref ref-type="bibr" rid="B15">Chambon et al., 2017</xref>). Performance was worst on NREM1. This makes sense given that NREM1 is the smallest class and can be similar to Awake and REM (<xref ref-type="bibr" rid="B34">Iber et al., 2007</xref>; <xref ref-type="bibr" rid="B76">Tsinalis et al., 2016a</xref>). NREM1 classification has often been relatively poor in previous studies (<xref ref-type="bibr" rid="B76">Tsinalis et al., 2016a</xref>; <xref ref-type="bibr" rid="B72">Supratak et al., 2017</xref>; <xref ref-type="bibr" rid="B14">Chambon et al., 2018</xref>; <xref ref-type="bibr" rid="B46">Michielli et al., 2019</xref>). Although Awake and NREM1 had similar numbers of samples, the classifier performed much better on Awake. Given that Awake EEG and EOG have features that are very different from those of NREM and that Awake EMG is different from REM EMG (<xref ref-type="bibr" rid="B34">Iber et al., 2007</xref>), it makes sense that it would be easier to classify Awake samples. Similar to previous studies (<xref ref-type="bibr" rid="B15">Chambon et al., 2017</xref>; <xref ref-type="bibr" rid="B72">Supratak et al., 2017</xref>), the precision and F1 score, but not recall, were highest for NREM2.</p>
</sec>
<sec id="S4.SS3">
<title>4.3. Global results</title>
<p>Across methods, EEG was most important for identifying Awake, NREM2, NREM3, and REM. In contrast, EOG played a greater role in the correct classification of NREM1 samples. EMG was not very important to the classification of any stage. This result is not atypical, as previous studies have shown that using EEG and EMG does not greatly improve classification performance for Awake, NREM2, and NREM3 relative to only using EEG (<xref ref-type="bibr" rid="B37">Kim and Choi, 2018</xref>). While global explanations for correct classification groups were similar across methods, explanations for incorrect groups tended to differ across methods.</p>
</sec>
<sec id="S4.SS4">
<title>4.4. Subject-level local ablation and LRP results over time</title>
<p>The two local approaches had similar results for the 2-h period of explanations that we output. In contrast to the global explanations, EOG was particularly important during Awake periods. This suggests that subject or subgroup-specific patterns of EOG activity exist within the Awake class that are obscured by global methods. It also supports existing findings that EEG alone did not discriminate between Awake, NREM1, and REM as effectively as EEG with EOG and EMG (<xref ref-type="bibr" rid="B30">Estrada et al., 2006</xref>). Additionally, previous studies have found that EOG is particularly important for identifying Awake (<xref ref-type="bibr" rid="B56">Pettersson et al., 2019</xref>) and can yield comparable classification performance to EEG (<xref ref-type="bibr" rid="B31">Ganesan and Jain, 2020</xref>). In contrast, EEG was important for discriminating NREM and REM samples, which makes sense given that NREM and REM EEG differ greatly (<xref ref-type="bibr" rid="B34">Iber et al., 2007</xref>). It is interesting that the subject had higher Awake EOG than EEG importance. Globally, EEG tended to be more important for Awake than EOG. Moreover, visualizing the results over time enabled us to obtain higher resolution insight into the classifier than global visualization. For example, EMG importance for the subject spiked for incorrectly classified samples, which suggests that EMG adversely affected model performance for some subjects.</p>
</sec>
<sec id="S4.SS5">
<title>4.5. Statistical analysis of effects of clinical and demographic variables upon local explanations</title>
<p>Interestingly, sex has relationships with more modality correct classification group pairs than either medication or age, which could indicate that subject sex had stronger effects on the patterns learned by the classifier than the other variables. This is potentially attributable to the imbalance of male and female subjects. Subject sex seemed to affect the NREM2 EEG patterns learned by the classifier. This reflects established sleep science. Namely, adult women can have greater slow-wave EEG activity in NREM sleep stages than men (<xref ref-type="bibr" rid="B49">Mourtazaev et al., 1995</xref>; <xref ref-type="bibr" rid="B19">Ehlers and Kupfer, 1997</xref>), and, in general, there are differences in the EEG activity of men and women (<xref ref-type="bibr" rid="B5">Armitage and Hoffmann, 2001</xref>; <xref ref-type="bibr" rid="B12">Bu&#x010D;kov&#x00E1; et al., 2020</xref>). While sex was associated with the correct classification of NREM1, sex may have adversely affected the EEG patterns learned by the classifier for Awake, NREM1, NREM2, and REM. Whereas the effects of sex on EEG was more associated with incorrect classification, both explainability methods indicated that sex likely affected the EOG patterns learned by the classifier for the correct classification of Awake, NREM2, NREM3, and REM. This highlights the possibility of EOG sex differences across most sleep stages. Our literature review has uncovered no studies on the effects of sex upon EOG in sleep, so our results could prompt future studies on this topic. Both methods indicated that sex affected the EMG patterns learned for incorrectly classified samples. Medication affected the EEG of Awake, NREM3, and REM similarly, with both methods. Previous studies have shown that benzodiazepines like temazepam (<xref ref-type="bibr" rid="B10">Bastien et al., 2003</xref>) and other medications (<xref ref-type="bibr" rid="B13">Chalon et al., 2005</xref>) can have significant effects on EEG sleep stages and that temazepam, in particular, can greatly affect REM (<xref ref-type="bibr" rid="B53">Pagel and Farnes, 2001</xref>). Other studies have shown similar effects in monkey EEG (<xref ref-type="bibr" rid="B7">Authier et al., 2014</xref>). Our results also showed that medication significantly affected the patterns learned for REM EOG. Interestingly, medication may have been related to the learning of EMG patterns that contributed to incorrect NREM classification. The inconsistent effects of medication upon EMG could fit with previous studies that purportedly analyzed EMG sleep data in monkeys but did not report any effects of medication (<xref ref-type="bibr" rid="B7">Authier et al., 2014</xref>). The effects of age on sleep are well characterized (<xref ref-type="bibr" rid="B49">Mourtazaev et al., 1995</xref>; <xref ref-type="bibr" rid="B19">Ehlers and Kupfer, 1997</xref>; <xref ref-type="bibr" rid="B11">Boselli et al., 1998</xref>; <xref ref-type="bibr" rid="B17">Chinoy Frey et al., 2014</xref>; <xref ref-type="bibr" rid="B44">Luca et al., 2015</xref>). In our study, age seemed to affect the EEG patterns learned for REM like in <xref ref-type="bibr" rid="B41">Landolt and Borb&#x00E9;ly (2001)</xref>. However, age was also related to the learning of EEG patterns for multiple incorrect classification groups. This suggests that the model did not fully learn to address the underlying effects of age upon EEG across sleep stages. Interestingly, age had inconsistent effects upon the EOG patterns learned by the classifier. Age seemed to affect EMG patterns for NREM1 and NREM2.</p>
</sec>
<sec id="S4.SS6">
<title>4.6. Limitations and next steps</title>
<p>Future studies might compare differences in importance across more subjects, which could help identify personalized sleep stage biomarkers (<xref ref-type="bibr" rid="B58">Porumb et al., 2020</xref>). For our line-related noise global ablation approach, we used a combination of a 50-Hz sinusoid and Gaussian noise. This approach provided a useful proof-of-concept and is viable for use in future studies. However, it only provides a simple simulation of line noise. Line noise can, in practice, have a more complex power spectral density around 50 Hz. In this study, we used a simple CNN classifier, which made the implementation of LRP straightforward. However, using a simple CNN classifier also contributed to classification performance that was high but below the state of the art. Future studies with advanced classifiers might use the analyses that we employed to assist with the discovery of biomarkers and formulation of novel hypotheses related to sleep and other domains. Our classifier was originally developed for EEG sleep stage classification. As such, the architecture may not be optimized for EOG and EMG feature extraction. While this does not adversely affect the quality of our explainability results, it prevents generalizable claims regarding the importance of one modality over another. Additionally, other GBFA methods could potentially replace LRP for multimodal explainability. Metrics like those presented in <xref ref-type="bibr" rid="B65">Samek et al. (2017a)</xref>, <xref ref-type="bibr" rid="B55">Petsiuk et al. (2018)</xref> could help rate the quality of each explainability method, and future studies might enhance LRP explanation quality by applying different relevance rules to different parts of a network (<xref ref-type="bibr" rid="B66">Samek et al., 2017b</xref>). Additionally, while our analysis of relationships between local explanations and clinical and demographic variables was insightful, future studies might perform a variety of other analyses on local explanations (<xref ref-type="bibr" rid="B74">Thoret et al., 2022</xref>). For example, they might cluster local explanations to identify subtypes of individuals or compute measures that quantify aspects of the temporal distribution of importance. Lastly, our dataset was only composed of data from 22 participants. As such, the generalizability of the conclusions that can be drawn from our analysis of the relationship between the local explanations and clinical and demographic variables is somewhat limited. Nevertheless, the analysis represents a novel approach for the domain and offers inspiration as a starting point for future studies in the field.</p>
</sec>
</sec>
<sec id="S5" sec-type="conclusion">
<title>5. Conclusion</title>
<p>In this study, we use sleep stage classification as a testbed for developing multimodal explainability methods. After training a classifier for multimodal sleep stage classification, we present a series of novel multimodal explainability methods. Up to this point, relatively few studies in the domain of multimodal classification have involved explainability, which is particularly concerning for clinical settings. Our global ablation method is uniquely adapted to electrophysiology classification. Our local ablation approach is the first local multimodal ablation method, and our GBFA approach offers an alternative to ablation that has not previously been used for modality importance. We find that EEG was most important to the identification of most sleep stages while EOG was most important to the identification of NREM1. We show how local methods can help identify differences in subject-level explanations that could potentially be used to identify personalized biomarkers in future studies. Importantly, we also developed a novel analysis approach and found that subject sex had more significant relationships with patterns learned by the classifier relative to other clinical and demographic variables. More broadly, the approach could help illuminate the effects of those variables upon different classes (e.g., sleep stages or disease conditions). Our study enhances the level of insight that can be obtained from the typically black-box models of the growing field of multimodal classification and has implications for personalized medicine and the eventual development of multimodal clinical classifiers.</p>
</sec>
<sec id="S6" sec-type="data-availability">
<title>Data availability statement</title>
<p>Publicly available datasets were analyzed in this study. This data can be found here: <ext-link ext-link-type="uri" xlink:href="https://www.physionet.org/content/sleep-edfx/1.0.0/">https://www.physionet.org/content/sleep-edfx/1.0.0/</ext-link>.</p>
</sec>
<sec id="S7" sec-type="author-contributions">
<title>Author contributions</title>
<p>CE helped with the conception of the manuscript, performed the analyses, wrote the manuscript, and edited the manuscript. MS helped with figure creation, writing, and editing the manuscript. RZ and DC helped perform analyses and edited the manuscript. MW and RM helped with conception of the manuscript and edited the manuscript. VC helped with the conception of the manuscript, edited the manuscript, and provided funding for the manuscript. All authors contributed to the article and approved the submitted version.</p>
</sec>
</body>
<back>
<sec id="S8" sec-type="funding-information">
<title>Funding</title>
<p>This work was funded by the NIH grant R01EB006841.</p>
</sec>
<ack><p>We thank those who collected the Sleep-EDF Database Expanded on PhysioNet.</p>
</ack>
<sec id="S9" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="S10" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec id="S11" sec-type="supplementary-material">
<title>Supplementary material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="https://www.frontiersin.org/articles/10.3389/fninf.2023.1123376/full#supplementary-material">https://www.frontiersin.org/articles/10.3389/fninf.2023.1123376/full#supplementary-material</ext-link></p>
<supplementary-material xlink:href="Data_Sheet_1.DOCX" id="DS1" mimetype="application/vnd.openxmlformats-officedocument.wordprocessingml.document" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Abadi</surname> <given-names>M.</given-names></name> <name><surname>Barham</surname> <given-names>P.</given-names></name> <name><surname>Chen</surname> <given-names>J.</given-names></name> <name><surname>Chen</surname> <given-names>Z.</given-names></name> <name><surname>Davis</surname> <given-names>A.</given-names></name> <name><surname>Dean</surname> <given-names>J.</given-names></name><etal/></person-group> (<year>2016</year>). &#x201C;<article-title>TensorFlow: A system for large-scale machine learning</article-title>,&#x201D; in <source><italic>Proceedings of the 12th USENIX Symposium on Operating Systems Design and Implementation</italic></source>, <publisher-loc>Savannah, GA</publisher-loc>.</citation></ref>
<ref id="B2"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Aboalayon</surname> <given-names>K.</given-names></name> <name><surname>Almuhammadi</surname> <given-names>W.</given-names></name> <name><surname>Faezipour</surname> <given-names>M.</given-names></name></person-group> (<year>2015</year>). &#x201C;<article-title>A comparison of different machine learning algorithms using single channel EEG signal for classifying human sleep stages</article-title>,&#x201D; in <source><italic>Proceedings of the 2015 IEEE Long Island Systems, Applications and Technology Conference, LISAT 2015</italic></source>, (<publisher-loc>Farmingdale, NY</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>1</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1109/LISAT.2015.7160185</pub-id></citation></ref>
<ref id="B3"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Alber</surname> <given-names>M.</given-names></name> <name><surname>Lapuschkin</surname> <given-names>S.</given-names></name> <name><surname>Seegerer</surname> <given-names>P.</given-names></name> <name><surname>H&#x00E4;gele</surname> <given-names>M.</given-names></name> <name><surname>Sch&#x00FC;tt</surname> <given-names>K.</given-names></name> <name><surname>Montavon</surname> <given-names>G.</given-names></name><etal/></person-group> (<year>2019</year>). <article-title>INNvestigate neural networks!</article-title> <source><italic>J. Mach. Learn. Res.</italic></source> <volume>20</volume> <fpage>1</fpage>&#x2013;<lpage>8</lpage>.</citation></ref>
<ref id="B4"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ancona</surname> <given-names>M.</given-names></name> <name><surname>Ceolini</surname> <given-names>E.</given-names></name> <name><surname>&#x00D6;ztireli</surname> <given-names>C.</given-names></name> <name><surname>Gross</surname> <given-names>M.</given-names></name></person-group> (<year>2018</year>). &#x201C;<article-title>Towards better understanding of gradient-based attribution methods for deep neural networks</article-title>,&#x201D; in <source><italic>Proceedings of the International Conference on Learning Representations</italic></source>, <publisher-loc>Z&#x00FC;rich</publisher-loc>. <pub-id pub-id-type="doi">10.1007/978-3-030-28954-6_9</pub-id></citation></ref>
<ref id="B5"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Armitage</surname> <given-names>R.</given-names></name> <name><surname>Hoffmann</surname> <given-names>R.</given-names></name></person-group> (<year>2001</year>). <article-title>Sleep EEG, depression and gender.</article-title> <source><italic>Sleep Med. Rev.</italic></source> <volume>5</volume> <fpage>237</fpage>&#x2013;<lpage>246</lpage>. <pub-id pub-id-type="doi">10.1053/smrv.2000.0144</pub-id> <pub-id pub-id-type="pmid">12530989</pub-id></citation></ref>
<ref id="B6"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Arras</surname> <given-names>L.</given-names></name> <name><surname>Montavon</surname> <given-names>G.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>K.</given-names></name> <name><surname>Samek</surname> <given-names>W.</given-names></name></person-group> (<year>2017</year>). &#x201C;<article-title>Explaining Recurrent Neural Network Predictions in Sentiment Analysis</article-title>,&#x201D; in <source><italic>Proceedings of the EMNLP 2017 &#x2013; 8th workshop on computational approaches to subjectivity, sentiment &#x0026; social media Analysis</italic></source> <publisher-loc>Copenhagen</publisher-loc>, <fpage>159</fpage>&#x2013;<lpage>168</lpage>. <pub-id pub-id-type="doi">10.18653/v1/W17-5221</pub-id> <pub-id pub-id-type="pmid">36568019</pub-id></citation></ref>
<ref id="B7"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Authier</surname> <given-names>S.</given-names></name> <name><surname>Bassett</surname> <given-names>L.</given-names></name> <name><surname>Pouliot</surname> <given-names>M.</given-names></name> <name><surname>Rachalski</surname> <given-names>A.</given-names></name> <name><surname>Troncy</surname> <given-names>E.</given-names></name> <name><surname>Paquette</surname> <given-names>D.</given-names></name><etal/></person-group> (<year>2014</year>). <article-title>Effects of amphetamine, diazepam and caffeine on polysomnography (EEG, EMG, EOG)-derived variables measured using telemetry in Cynomolgus monkeys.</article-title> <source><italic>J. Pharmacol. Toxicol. Methods</italic></source> <volume>70</volume> <fpage>86</fpage>&#x2013;<lpage>93</lpage>. <pub-id pub-id-type="doi">10.1016/j.vascn.2014.05.003</pub-id> <pub-id pub-id-type="pmid">24878255</pub-id></citation></ref>
<ref id="B8"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bach</surname> <given-names>S.</given-names></name> <name><surname>Binder</surname> <given-names>A.</given-names></name> <name><surname>Montavon</surname> <given-names>G.</given-names></name> <name><surname>Klauschen</surname> <given-names>F.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>K.</given-names></name> <name><surname>Samek</surname> <given-names>W.</given-names></name></person-group> (<year>2015</year>). <article-title>On pixel-wise explanations for non-linear classifier decisions by layer-wise relevance propagation.</article-title> <source><italic>PLoS One</italic></source> <volume>10</volume>:<issue>e0130140</issue>. <pub-id pub-id-type="doi">10.1371/journal.pone.0130140</pub-id> <pub-id pub-id-type="pmid">26161953</pub-id></citation></ref>
<ref id="B9"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Barnes</surname> <given-names>L.</given-names></name> <name><surname>Lee</surname> <given-names>K.</given-names></name> <name><surname>Kempa-Liehr</surname> <given-names>A.</given-names></name> <name><surname>Hallum</surname> <given-names>L.</given-names></name></person-group> (<year>2021</year>). <article-title>Detection of sleep apnea from single-channel electroencephalogram (EEG) using an explainable convolutional neural network.</article-title> <source><italic>PLoS One</italic></source> <volume>17</volume>:<issue>e0272167</issue>. <pub-id pub-id-type="doi">10.1371/journal.pone.0272167</pub-id> <pub-id pub-id-type="pmid">36099242</pub-id></citation></ref>
<ref id="B10"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bastien</surname> <given-names>C.</given-names></name> <name><surname>LeBlanc</surname> <given-names>M.</given-names></name> <name><surname>Carrier</surname> <given-names>J.</given-names></name> <name><surname>Morin</surname> <given-names>C.</given-names></name></person-group> (<year>2003</year>). <article-title>Sleep EEG power spectra, insomnia, and chronic use of benzodiazepines.</article-title> <source><italic>Sleep</italic></source> <volume>26</volume> <fpage>313</fpage>&#x2013;<lpage>317</lpage>. <pub-id pub-id-type="doi">10.1093/sleep/26.3.313</pub-id> <pub-id pub-id-type="pmid">12749551</pub-id></citation></ref>
<ref id="B11"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Boselli</surname> <given-names>M.</given-names></name> <name><surname>Parrino</surname> <given-names>L.</given-names></name> <name><surname>Smerieri</surname> <given-names>A.</given-names></name> <name><surname>Terzano</surname> <given-names>M.</given-names></name></person-group> (<year>1998</year>). <article-title>Effect of age on EEG arousals in normal sleep.</article-title> <source><italic>Sleep</italic></source> <volume>21</volume> <fpage>361</fpage>&#x2013;<lpage>367</lpage>.</citation></ref>
<ref id="B12"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bu&#x010D;kov&#x00E1;</surname> <given-names>B.</given-names></name> <name><surname>Brunovsk&#x00FD;</surname> <given-names>M.</given-names></name> <name><surname>Bare&#x0161;</surname> <given-names>M.</given-names></name> <name><surname>Hlinka</surname> <given-names>J.</given-names></name></person-group> (<year>2020</year>). <article-title>Predicting sex from EEG: Validity and generalizability of deep-learning-based interpretable classifier.</article-title> <source><italic>Front. Neurosci.</italic></source> <volume>14</volume>:<issue>589303</issue>. <pub-id pub-id-type="doi">10.3389/fnins.2020.589303</pub-id> <pub-id pub-id-type="pmid">33192274</pub-id></citation></ref>
<ref id="B13"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chalon</surname> <given-names>S.</given-names></name> <name><surname>Pereira</surname> <given-names>A.</given-names></name> <name><surname>Lainey</surname> <given-names>E.</given-names></name> <name><surname>Vandenhende</surname> <given-names>F.</given-names></name> <name><surname>Watkin</surname> <given-names>J.</given-names></name> <name><surname>Staner</surname> <given-names>L.</given-names></name><etal/></person-group> (<year>2005</year>). <article-title>Comparative effects of duloxetine and desipramine on sleep EEG in healthy subjects.</article-title> <source><italic>Psychopharmacology</italic></source> <volume>177</volume> <fpage>357</fpage>&#x2013;<lpage>365</lpage>. <pub-id pub-id-type="doi">10.1007/s00213-004-1961-0</pub-id> <pub-id pub-id-type="pmid">15290000</pub-id></citation></ref>
<ref id="B14"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chambon</surname> <given-names>S.</given-names></name> <name><surname>Galtier</surname> <given-names>M.</given-names></name> <name><surname>Arnal</surname> <given-names>P.</given-names></name> <name><surname>Wainrib</surname> <given-names>G.</given-names></name> <name><surname>Gramfort</surname> <given-names>A.</given-names></name></person-group> (<year>2018</year>). <article-title>A deep learning architecture for temporal sleep stage classification using multivariate and multimodal time series.</article-title> <source><italic>IEEE Trans. Neural. Syst. Rehabil. Eng.</italic></source> <volume>26</volume> <fpage>758</fpage>&#x2013;<lpage>769</lpage>.</citation></ref>
<ref id="B15"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chambon</surname> <given-names>S.</given-names></name> <name><surname>Galtier</surname> <given-names>M.</given-names></name> <name><surname>Arnal</surname> <given-names>P.</given-names></name> <name><surname>Wainrib</surname> <given-names>G.</given-names></name> <name><surname>Gramfort</surname> <given-names>A.</given-names></name></person-group> (<year>2017</year>). <article-title>A deep learning architecture for temporal sleep stage classification using multivariate and multimodal time series.</article-title> <source><italic>IEEE Trans. Neural Syst. Rehabil. Eng.</italic></source> <volume>26</volume> <fpage>758</fpage>&#x2013;<lpage>769</lpage>.</citation></ref>
<ref id="B16"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Gong</surname> <given-names>C.</given-names></name> <name><surname>Hao</surname> <given-names>H.</given-names></name> <name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Xu</surname> <given-names>S.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name><etal/></person-group> (<year>2019</year>). <article-title>Automatic sleep stage classification based on subthalamic local field potentials.</article-title> <source><italic>IEEE Trans. Neural. Syst. Rehabil. Eng.</italic></source> <volume>27</volume> <fpage>118</fpage>&#x2013;<lpage>128</lpage>. <pub-id pub-id-type="doi">10.1109/TNSRE.2018.2890272</pub-id> <pub-id pub-id-type="pmid">30605104</pub-id></citation></ref>
<ref id="B17"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chinoy Frey</surname> <given-names>D.</given-names></name> <name><surname>Kaslovsky</surname> <given-names>D.</given-names></name> <name><surname>Meyer</surname> <given-names>F.</given-names></name> <name><surname>Wright</surname> <given-names>K.</given-names></name></person-group> (<year>2014</year>). <article-title>Age-related changes in slow wave activity rise time and NREM sleep EEG with and without zolpidem in healthy young and older adults.</article-title> <source><italic>Sleep Med.</italic></source> <volume>15</volume> <fpage>1037</fpage>&#x2013;<lpage>1045</lpage>. <pub-id pub-id-type="doi">10.1016/j.sleep.2014.05.007</pub-id> <pub-id pub-id-type="pmid">24980066</pub-id></citation></ref>
<ref id="B18"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chollet</surname> <given-names>F.</given-names></name></person-group> (<year>2015</year>). <source><italic>Keras.</italic></source> Available from: <ext-link ext-link-type="uri" xlink:href="https://github.com/fchollet/keras">https://github.com/fchollet/keras</ext-link>. <comment>(accessed December 13, 2022)</comment>.</citation></ref>
<ref id="B19"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ehlers</surname> <given-names>C.</given-names></name> <name><surname>Kupfer</surname> <given-names>D.</given-names></name></person-group> (<year>1997</year>). <article-title>Slow-wave sleep: Do young adult men and women age differently?</article-title> <source><italic>J. Sleep Res.</italic></source> <volume>6</volume> <fpage>211</fpage>&#x2013;<lpage>215</lpage>. <pub-id pub-id-type="doi">10.1046/j.1365-2869.1997.00041.x</pub-id> <pub-id pub-id-type="pmid">9358400</pub-id></citation></ref>
<ref id="B20"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Eldele</surname> <given-names>E.</given-names></name> <name><surname>Chen</surname> <given-names>Z.</given-names></name> <name><surname>Liu</surname> <given-names>C.</given-names></name> <name><surname>Wu</surname> <given-names>M.</given-names></name> <name><surname>Kwoh</surname> <given-names>C.</given-names></name> <name><surname>Li</surname> <given-names>X.</given-names></name><etal/></person-group> (<year>2021</year>). <article-title>An attention-based deep learning approach for sleep stage classification with single-channel EEG.</article-title> <source><italic>IEEE Trans. Neural. Syst. Rehabil. Eng.</italic></source> <volume>29</volume> <fpage>809</fpage>&#x2013;<lpage>818</lpage>. <pub-id pub-id-type="doi">10.1109/TNSRE.2021.3076234</pub-id> <pub-id pub-id-type="pmid">33909566</pub-id></citation></ref>
<ref id="B21"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Carbajal</surname> <given-names>D.</given-names></name> <name><surname>Zhang</surname> <given-names>R.</given-names></name> <name><surname>Miller</surname> <given-names>R.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name> <name><surname>Wang</surname> <given-names>M.</given-names></name></person-group> (<year>2021a</year>). <article-title>An explainable deep learning approach for multimodal electrophysiology classification.</article-title> <source><italic>bioRxiv</italic></source> [<comment>Preprint</comment>]. <pub-id pub-id-type="doi">10.1101/2021.05.12.443594</pub-id></citation></ref>
<ref id="B22"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Carbajal</surname> <given-names>D.</given-names></name> <name><surname>Zhang</surname> <given-names>R.</given-names></name> <name><surname>Sendi</surname> <given-names>M.</given-names></name> <name><surname>Miller</surname> <given-names>R.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name><etal/></person-group> (<year>2021c</year>). &#x201C;<article-title>A novel local ablation approach for explaining multimodal classifiers</article-title>,&#x201D; in <source><italic>Proceedings of the 2021 IEEE 21st international conference on bioinformatics and bioengineering (BIBE)</italic></source>, <publisher-loc>Kragujevac</publisher-loc>. <pub-id pub-id-type="doi">10.1109/BIBE52308.2021.9635541</pub-id></citation></ref>
<ref id="B23"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Miller</surname> <given-names>R.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name></person-group> (<year>2021f</year>). &#x201C;<article-title>A novel local explainability approach for spectral insight into raw EEG-based deep learning classifiers</article-title>,&#x201D; in <source><italic>Proceedings of the 21st IEEE international conference on bioinformatics and bioengineering</italic></source>, <publisher-loc>Kragujevac</publisher-loc>. <pub-id pub-id-type="doi">10.1109/BIBE52308.2021.9635243</pub-id></citation></ref>
<ref id="B24"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Miller</surname> <given-names>R.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name></person-group> (<year>2021g</year>). <article-title>A gradient-based spectral explainability method for EEG deep learning classifiers.</article-title> <source><italic>bioRxiv</italic></source> [<comment>Preprint</comment>]. <pub-id pub-id-type="doi">10.1101/2021.07.14.452360</pub-id></citation></ref>
<ref id="B25"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Miller</surname> <given-names>R.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name></person-group> (<year>2021i</year>). <article-title>A model visualization-based approach for insight into waveforms and spectra learned by CNNs.</article-title> <source><italic>bioRxiv</italic></source> [<comment>Preprint</comment>]. <pub-id pub-id-type="doi">10.1101/2021.12.16.473028</pub-id></citation></ref>
<ref id="B26"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Miller</surname> <given-names>R.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name> <name><surname>Wang</surname> <given-names>M.</given-names></name></person-group> (<year>2021d</year>). &#x201C;<article-title>A gradient-based approach for explaining multimodal deep learning classifiers</article-title>,&#x201D; in <source><italic>Proceedings of the 2021 IEEE 21st international conference on bioinformatics and bioengineering (BIBE)</italic></source>, (<publisher-loc>Kragujevac</publisher-loc>: <publisher-name>IEEE</publisher-name>). <pub-id pub-id-type="doi">10.1109/BIBE52308.2021.9635460</pub-id></citation></ref>
<ref id="B27"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Sendi</surname> <given-names>M.</given-names></name> <name><surname>Miller</surname> <given-names>R.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name></person-group> (<year>2021h</year>). &#x201C;<article-title>A novel activation maximization-based approach for insight into electrophysiology classifiers</article-title>,&#x201D; in <source><italic>Proceedings of the 2021 IEEE international conference on bioinformatics and biomedicine (BIBM)</italic></source>, <publisher-loc>Houston, TX</publisher-loc>. <pub-id pub-id-type="doi">10.1109/BIBM52615.2021.9669593</pub-id></citation></ref>
<ref id="B28"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Sendi</surname> <given-names>M.</given-names></name> <name><surname>Willie</surname> <given-names>J.</given-names></name> <name><surname>Mahmoudi</surname> <given-names>B.</given-names></name></person-group> (<year>2021e</year>). &#x201C;<article-title>Hierarchical neural network with layer-wise relevance propagation for interpretable multiclass neural state classification</article-title>,&#x201D; in <source><italic>Proceedings of the 10th international IEEE/EMBS conference on neural engineering (NER)</italic></source>, <fpage>18</fpage>&#x2013;<lpage>21</lpage>. <pub-id pub-id-type="doi">10.1109/NER49283.2021.9441217</pub-id></citation></ref>
<ref id="B29"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ellis</surname> <given-names>C.</given-names></name> <name><surname>Zhang</surname> <given-names>R.</given-names></name> <name><surname>Carbajal</surname> <given-names>D.</given-names></name> <name><surname>Miller</surname> <given-names>R.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name> <name><surname>Wang</surname> <given-names>M.</given-names></name></person-group> (<year>2021b</year>). &#x201C;<article-title>Explainable Sleep Stage Classification with Multimodal Electrophysiology Time-series</article-title>,&#x201D; in <source><italic>Proceedings of the 2021 43rd Annual International Conference of the IEEE Engineering in Medicine &#x0026; Biology Society (EMBC)</italic></source>, <publisher-loc>Mexico</publisher-loc>. <pub-id pub-id-type="doi">10.1109/EMBC46164.2021.9630506</pub-id> <pub-id pub-id-type="pmid">34891757</pub-id></citation></ref>
<ref id="B30"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Estrada</surname> <given-names>E.</given-names></name> <name><surname>Nazeran</surname> <given-names>H.</given-names></name> <name><surname>Barragan</surname> <given-names>J.</given-names></name> <name><surname>Burk</surname> <given-names>J.</given-names></name> <name><surname>Lucas</surname> <given-names>E.</given-names></name> <name><surname>Behbehani</surname> <given-names>K.</given-names></name></person-group> (<year>2006</year>). <article-title>EOG and EMG: Two important switches in automatic sleep stage classification.</article-title> <source><italic>Conf. Proc. IEEE Eng. Med. Biol. Soc.</italic></source> <volume>2006</volume> <fpage>2458</fpage>&#x2013;<lpage>2461</lpage>. <pub-id pub-id-type="doi">10.1109/IEMBS.2006.260075</pub-id> <pub-id pub-id-type="pmid">17946514</pub-id></citation></ref>
<ref id="B31"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ganesan</surname> <given-names>R.</given-names></name> <name><surname>Jain</surname> <given-names>R.</given-names></name></person-group> (<year>2020</year>). &#x201C;<article-title>Binary state prediction of sleep or wakefulness using EEG and EOG features</article-title>,&#x201D; in <source><italic>Proceedings of the 2020 IEEE 17th India council international conference (INDICON)</italic></source>, <publisher-loc>New Delhi</publisher-loc>. <pub-id pub-id-type="doi">10.1109/INDICON49873.2020.9342272</pub-id></citation></ref>
<ref id="B32"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Goldberger</surname> <given-names>A.</given-names></name> <name><surname>Amaral</surname> <given-names>L.</given-names></name> <name><surname>Glass</surname> <given-names>L.</given-names></name> <name><surname>Hausdorff</surname> <given-names>J.</given-names></name> <name><surname>Ivanov</surname> <given-names>P.</given-names></name> <name><surname>Mark</surname> <given-names>R.</given-names></name><etal/></person-group> (<year>2000</year>). <article-title>PhysioBank, PhysioToolkit, and PhysioNet: Components of a new research resource for complex physiologic signals.</article-title> <source><italic>Circulation</italic></source> <volume>101</volume> <fpage>e215</fpage>&#x2013;<lpage>e220</lpage>. <pub-id pub-id-type="doi">10.1161/01.CIR.101.23.e215</pub-id> <pub-id pub-id-type="pmid">10851218</pub-id></citation></ref>
<ref id="B33"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Griffin</surname> <given-names>C.</given-names></name> <name><surname>Kaye</surname> <given-names>A.</given-names></name> <name><surname>Rivera Bueno</surname> <given-names>F.</given-names></name> <name><surname>Kaye</surname> <given-names>A.</given-names></name></person-group> (<year>2013</year>). <article-title>Benzodiazepine pharmacology and central nervous system-mediated effects.</article-title> <source><italic>Ochsner J.</italic></source> <volume>13</volume> <fpage>214</fpage>&#x2013;<lpage>223</lpage>.</citation></ref>
<ref id="B34"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Iber</surname> <given-names>C.</given-names></name> <name><surname>Ancoli-Israel</surname> <given-names>S.</given-names></name> <name><surname>Chesson</surname> <given-names>A.</given-names></name> <name><surname>Quan</surname> <given-names>S.</given-names></name></person-group> (<year>2007</year>). <source><italic>The AASM manual for scoring of sleep and associated events: Rules, terminology, and technical specifications</italic></source>, <publisher-loc>Westchester, IL</publisher-loc>: <publisher-name>American Academy of Sleep Medicine</publisher-name>.</citation></ref>
<ref id="B35"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kemp</surname> <given-names>B.</given-names></name> <name><surname>Zwinderman</surname> <given-names>A.</given-names></name> <name><surname>Tuk</surname> <given-names>B.</given-names></name> <name><surname>Kamphuisen</surname> <given-names>H.</given-names></name> <name><surname>Oberye</surname> <given-names>J.</given-names></name></person-group> (<year>2000</year>). <article-title>Analysis of a sleep-dependent neuronal feedback loop: The slow-wave microcontinuity of the EEG.</article-title> <source><italic>IEEE Trans. Biomed. Eng.</italic></source> <volume>47</volume> <fpage>1185</fpage>&#x2013;<lpage>1194</lpage>. <pub-id pub-id-type="doi">10.1109/10.867928</pub-id> <pub-id pub-id-type="pmid">11008419</pub-id></citation></ref>
<ref id="B36"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Khalighi</surname> <given-names>S.</given-names></name> <name><surname>Sousa</surname> <given-names>T.</given-names></name> <name><surname>Santos</surname> <given-names>J.</given-names></name> <name><surname>Nunes</surname> <given-names>U.</given-names></name></person-group> (<year>2016</year>). <article-title>ISRUC-Sleep: A comprehensive public dataset for sleep researchers.</article-title> <source><italic>Comput. Methods Programs Biomed.</italic></source> <volume>124</volume> <fpage>180</fpage>&#x2013;<lpage>192</lpage>. <pub-id pub-id-type="doi">10.1016/j.cmpb.2015.10.013</pub-id> <pub-id pub-id-type="pmid">26589468</pub-id></citation></ref>
<ref id="B37"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kim</surname> <given-names>H.</given-names></name> <name><surname>Choi</surname> <given-names>S.</given-names></name></person-group> (<year>2018</year>). <article-title>&#x201C;Automatic Sleep Stage Classification Using EEG and EMG Signal,&#x201D;</article-title> in <source><italic>Proceedings of the 2018 tenth international conference on ubiquitous and future networks (ICUFN)</italic></source>, <publisher-loc>Prague</publisher-loc>, <fpage>207</fpage>&#x2013;<lpage>212</lpage>.</citation></ref>
<ref id="B38"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kingma</surname> <given-names>D.</given-names></name> <name><surname>Ba</surname> <given-names>J.</given-names></name></person-group> (<year>2015</year>). &#x201C;<article-title>Adam: A method for stochastic optimization</article-title>,&#x201D; in <source><italic>Proceedings of the 3rd international conference on learning representations (ICLR)</italic></source>, <publisher-loc>San Diego, CA</publisher-loc>.</citation></ref>
<ref id="B39"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kwon</surname> <given-names>Y.</given-names></name> <name><surname>Shin</surname> <given-names>S.</given-names></name> <name><surname>Kim</surname> <given-names>S.</given-names></name></person-group> (<year>2018</year>). <article-title>Electroencephalography based fusion two-dimensional (2D)-convolution neural networks (CNN) model for emotion recognition system.</article-title> <source><italic>Sensors</italic></source> <volume>18</volume>:<issue>1383</issue>. <pub-id pub-id-type="doi">10.3390/s18051383</pub-id> <pub-id pub-id-type="pmid">29710869</pub-id></citation></ref>
<ref id="B40"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lajnef</surname> <given-names>T.</given-names></name> <name><surname>Chaibi</surname> <given-names>S.</given-names></name> <name><surname>Ruby</surname> <given-names>P.</given-names></name> <name><surname>Aguera</surname> <given-names>P.</given-names></name> <name><surname>Eichenlaub</surname> <given-names>J.</given-names></name> <name><surname>Samet</surname> <given-names>M.</given-names></name><etal/></person-group> (<year>2015</year>). <article-title>Learning machines and sleeping brains: Automatic sleep stage classification using decision-tree multi-class support vector machines.</article-title> <source><italic>J. Neurosci. Methods</italic></source> <volume>250</volume> <fpage>94</fpage>&#x2013;<lpage>105</lpage>. <pub-id pub-id-type="doi">10.1016/j.jneumeth.2015.01.022</pub-id> <pub-id pub-id-type="pmid">25629798</pub-id></citation></ref>
<ref id="B41"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Landolt</surname> <given-names>H.</given-names></name> <name><surname>Borb&#x00E9;ly</surname> <given-names>A.</given-names></name></person-group> (<year>2001</year>). <article-title>Age-dependent changes in sleep EEG topography.</article-title> <source><italic>Clin. Neurophysiol.</italic></source> <volume>112</volume> <fpage>369</fpage>&#x2013;<lpage>377</lpage>. <pub-id pub-id-type="doi">10.1016/S1388-2457(00)00542-3</pub-id> <pub-id pub-id-type="pmid">11165543</pub-id></citation></ref>
<ref id="B42"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>Y.</given-names></name> <name><surname>Yang</surname> <given-names>X.</given-names></name> <name><surname>Zhi</surname> <given-names>X.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Cao</surname> <given-names>Z.</given-names></name></person-group> (<year>2021</year>). <source><italic>Automatic sleep stage classification based on two-channel EOG and one-channel EMG.</italic></source> Available online at: <ext-link ext-link-type="uri" xlink:href="https://www.researchsquare.com/article/rs-491468/latest?utm_source=researcher_app&#x0026;utm_medium=referral&#x0026;utm_campaign=RESR_MRKT_Researcher_inbound">https://www.researchsquare.com/article/rs-491468/latest?utm_source=researcher_app&#x0026;utm_medium=referral&#x0026;utm_campaign=RESR_MRKT_Researcher_inbound</ext-link> <comment>(accessed December 13, 2022)</comment>.</citation></ref>
<ref id="B43"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lin</surname> <given-names>J.</given-names></name> <name><surname>Pan</surname> <given-names>S.</given-names></name> <name><surname>Lee</surname> <given-names>C.</given-names></name> <name><surname>Oviatt</surname> <given-names>S.</given-names></name></person-group> (<year>2019</year>). &#x201C;<article-title>An explainable deep fusion network for affect recognition using physiological signals</article-title>,&#x201D; in <source><italic>Proceedings of the 28th ACM international conference on information and knowledge management</italic></source>, (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>ACM</publisher-name>). <pub-id pub-id-type="doi">10.1145/3357384.3358160</pub-id></citation></ref>
<ref id="B44"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Luca</surname> <given-names>G.</given-names></name> <name><surname>Haba Rubio</surname> <given-names>J.</given-names></name> <name><surname>Andries</surname> <given-names>D.</given-names></name> <name><surname>Tobback</surname> <given-names>N.</given-names></name> <name><surname>Vollenweider</surname> <given-names>P.</given-names></name> <name><surname>Waeber</surname> <given-names>G.</given-names></name><etal/></person-group> (<year>2015</year>). <article-title>Age and gender variations of sleep in subjects without sleep disorders.</article-title> <source><italic>Ann. Med.</italic></source> <volume>47</volume> <fpage>482</fpage>&#x2013;<lpage>491</lpage>. <pub-id pub-id-type="doi">10.3109/07853890.2015.1074271</pub-id> <pub-id pub-id-type="pmid">26224201</pub-id></citation></ref>
<ref id="B45"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mellem</surname> <given-names>M.</given-names></name> <name><surname>Liu</surname> <given-names>Y.</given-names></name> <name><surname>Gonzalez</surname> <given-names>H.</given-names></name> <name><surname>Kollada</surname> <given-names>M.</given-names></name> <name><surname>Martin</surname> <given-names>W.</given-names></name> <name><surname>Ahammad</surname> <given-names>P.</given-names></name></person-group> (<year>2020</year>). <article-title>Machine learning models identify multimodal measurements highly predictive of transdiagnostic symptom severity for mood, anhedonia, and anxiety.</article-title> <source><italic>Biol. Psychiatry Cogn. Neurosci. Neuroimaging</italic></source> <volume>5</volume> <fpage>56</fpage>&#x2013;<lpage>67</lpage>. <pub-id pub-id-type="doi">10.1016/j.bpsc.2019.07.007</pub-id> <pub-id pub-id-type="pmid">31543457</pub-id></citation></ref>
<ref id="B46"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Michielli</surname> <given-names>N.</given-names></name> <name><surname>Acharya</surname> <given-names>U.</given-names></name> <name><surname>Molinari</surname> <given-names>F.</given-names></name></person-group> (<year>2019</year>). <article-title>Cascaded LSTM recurrent neural network for automated sleep stage classification using single-channel EEG signals.</article-title> <source><italic>Comput. Biol. Med.</italic></source> <volume>106</volume> <fpage>71</fpage>&#x2013;<lpage>81</lpage>. <pub-id pub-id-type="doi">10.1016/j.compbiomed.2019.01.013</pub-id> <pub-id pub-id-type="pmid">30685634</pub-id></citation></ref>
<ref id="B47"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Molnar</surname> <given-names>C.</given-names></name></person-group> (<year>2018</year>). <source><italic>Interpretable machine learning. A guide for making black box models explainable.</italic></source> Available online at: <ext-link ext-link-type="uri" xlink:href="http://leanpub.com/interpretable-machine-learning">http://leanpub.com/interpretable-machine-learning</ext-link> <comment>(accessed August 08, 2018)</comment>.</citation></ref>
<ref id="B48"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Montavon</surname> <given-names>G.</given-names></name> <name><surname>Samek</surname> <given-names>W.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>K.</given-names></name></person-group> (<year>2018</year>). <article-title>Methods for interpreting and understanding deep neural networks.</article-title> <source><italic>Digit. Signal. Process A Rev. J.</italic></source> <volume>73</volume> <fpage>1</fpage>&#x2013;<lpage>15</lpage>. <pub-id pub-id-type="doi">10.1016/j.dsp.2017.10.011</pub-id></citation></ref>
<ref id="B49"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mourtazaev</surname> <given-names>M.</given-names></name> <name><surname>Kemp</surname> <given-names>B.</given-names></name> <name><surname>Zwinderman</surname> <given-names>A.</given-names></name> <name><surname>Kamphuisen</surname> <given-names>H.</given-names></name></person-group> (<year>1995</year>). <article-title>Age and gender affect different characteristics of slow waves in the sleep EEG.</article-title> <source><italic>Sleep</italic></source> <volume>18</volume> <fpage>557</fpage>&#x2013;<lpage>564</lpage>. <pub-id pub-id-type="doi">10.1093/sleep/18.7.557</pub-id> <pub-id pub-id-type="pmid">8552926</pub-id></citation></ref>
<ref id="B50"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mousavi</surname> <given-names>S.</given-names></name> <name><surname>Afghah</surname> <given-names>F.</given-names></name> <name><surname>Rajendra Acharya</surname> <given-names>U.</given-names></name></person-group> (<year>2019</year>). <article-title>SleepEEGNet: Automated sleep stage scoring with sequence to sequence deep learning approach.</article-title> <source><italic>PLoS One</italic></source> <volume>14</volume>:<issue>e0216456</issue>. <pub-id pub-id-type="doi">10.1371/journal.pone.0216456</pub-id> <pub-id pub-id-type="pmid">31063501</pub-id></citation></ref>
<ref id="B51"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Nahmias</surname> <given-names>D.</given-names></name> <name><surname>Kontson</surname> <given-names>K.</given-names></name></person-group> (<year>2020</year>). &#x201C;<article-title>Easy perturbation EEG algorithm for spectral importance (easyPEASI): A simple method to identify important spectral features of EEG in deep learning models</article-title>,&#x201D; in <source><italic>Proceedings of the 26th ACM SIGKDD International Conference on Knowledge Discovery &#x0026; Data Mining</italic></source>, (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>ACM</publisher-name>), <fpage>2398</fpage>&#x2013;<lpage>2406</lpage>. <pub-id pub-id-type="doi">10.1145/3394486.3403289</pub-id></citation></ref>
<ref id="B52"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Niroshana</surname> <given-names>S.</given-names></name> <name><surname>Zhu</surname> <given-names>X.</given-names></name> <name><surname>Chen</surname> <given-names>Y.</given-names></name> <name><surname>Chen</surname> <given-names>W.</given-names></name></person-group> (<year>2019</year>). &#x201C;<article-title>Sleep stage classification based on EEG, EOG, and CNN-GRU deep learning model</article-title>,&#x201D; in <source><italic>Proceedings of the 2019 IEEE 10th international conference on awareness science and technology (iCAST)</italic></source>, <publisher-loc>Morioka</publisher-loc>, <fpage>1</fpage>&#x2013;<lpage>7</lpage>.</citation></ref>
<ref id="B53"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pagel</surname> <given-names>J.</given-names></name> <name><surname>Farnes</surname> <given-names>B.</given-names></name></person-group> (<year>2001</year>). <article-title>Medications for the treatment of sleep disorders: An overview.</article-title> <source><italic>Prim. Care Companion J. Clin. Psychiatry</italic></source> <volume>3</volume> <fpage>118</fpage>&#x2013;<lpage>125</lpage>. <pub-id pub-id-type="doi">10.4088/PCC.v03n0303</pub-id> <pub-id pub-id-type="pmid">15014609</pub-id></citation></ref>
<ref id="B54"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pathak</surname> <given-names>S.</given-names></name> <name><surname>Lu</surname> <given-names>C.</given-names></name> <name><surname>Nagaraj</surname> <given-names>S.</given-names></name> <name><surname>van Putten</surname> <given-names>M.</given-names></name> <name><surname>Seifert</surname> <given-names>C.</given-names></name></person-group> (<year>2021</year>). <article-title>STQS: Interpretable multi-modal spatial-temporal-sequential model for automatic sleep scoring.</article-title> <source><italic>Artif. Intell. Med.</italic></source> <volume>114</volume>:<issue>102038</issue>. <pub-id pub-id-type="doi">10.1016/j.artmed.2021.102038</pub-id> <pub-id pub-id-type="pmid">33875157</pub-id></citation></ref>
<ref id="B55"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Petsiuk</surname> <given-names>V.</given-names></name> <name><surname>Das</surname> <given-names>A.</given-names></name> <name><surname>Saenko</surname> <given-names>K.</given-names></name></person-group> (<year>2018</year>). &#x201C;<article-title>RisE: Randomized input sampling for explanation of black-box models</article-title>,&#x201D; in <source><italic>Proceedings of the British machine vision conference 2018</italic></source>, <publisher-loc>Cardiff</publisher-loc>.</citation></ref>
<ref id="B56"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pettersson</surname> <given-names>K.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>K.</given-names></name> <name><surname>Tiet&#x00E4;v&#x00E4;inen</surname> <given-names>A.</given-names></name> <name><surname>Gould</surname> <given-names>K.</given-names></name> <name><surname>H&#x00E6;ggstr&#x00F6;m</surname> <given-names>E.</given-names></name></person-group> (<year>2019</year>). <article-title>Saccadic eye movements estimate prolonged time awake.</article-title> <source><italic>J. Sleep Res.</italic></source> <volume>28</volume> <fpage>1</fpage>&#x2013;<lpage>13</lpage>. <pub-id pub-id-type="doi">10.1111/jsr.12755</pub-id> <pub-id pub-id-type="pmid">30133045</pub-id></citation></ref>
<ref id="B57"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Phan</surname> <given-names>H.</given-names></name> <name><surname>Andreotti</surname> <given-names>F.</given-names></name> <name><surname>Cooray</surname> <given-names>N.</given-names></name> <name><surname>Chen</surname> <given-names>O.</given-names></name> <name><surname>De Vos</surname> <given-names>M.</given-names></name></person-group> (<year>2019</year>). <article-title>Joint classification and prediction CNN framework for automatic sleep stage classification.</article-title> <source><italic>IEEE Trans. Biomed. Eng.</italic></source> <volume>66</volume> <fpage>1285</fpage>&#x2013;<lpage>1296</lpage>. <pub-id pub-id-type="doi">10.1109/TBME.2018.2872652</pub-id> <pub-id pub-id-type="pmid">30346277</pub-id></citation></ref>
<ref id="B58"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Porumb</surname> <given-names>M.</given-names></name> <name><surname>Stranges</surname> <given-names>S.</given-names></name> <name><surname>Pescap&#x00E8;</surname> <given-names>A.</given-names></name> <name><surname>Pecchia</surname> <given-names>L.</given-names></name></person-group> (<year>2020</year>). <article-title>Precision medicine and artificial intelligence: A pilot study on deep learning for hypoglycemic events detection based on ECG.</article-title> <source><italic>Sci Rep.</italic></source> <volume>10</volume> <fpage>1</fpage>&#x2013;<lpage>16</lpage>. <pub-id pub-id-type="doi">10.1038/s41598-019-56927-5</pub-id> <pub-id pub-id-type="pmid">31932608</pub-id></citation></ref>
<ref id="B59"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Quan</surname> <given-names>S.</given-names></name> <name><surname>Howard</surname> <given-names>B. V.</given-names></name> <name><surname>Iber</surname> <given-names>C.</given-names></name> <name><surname>Kiley</surname> <given-names>J.</given-names></name> <name><surname>Nieto</surname> <given-names>F.</given-names></name> <name><surname>O&#x2019;Connor</surname> <given-names>G.</given-names></name><etal/></person-group> (<year>1997</year>). <article-title>The sleep heart health study: Design, rationale, and methods.</article-title> <source><italic>Sleep</italic></source> <volume>20</volume> <fpage>1077</fpage>&#x2013;<lpage>1085</lpage>.</citation></ref>
<ref id="B60"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rahman</surname> <given-names>M.</given-names></name> <name><surname>Bhuiyan</surname> <given-names>M.</given-names></name> <name><surname>Hassan</surname> <given-names>A.</given-names></name></person-group> (<year>2018</year>). <article-title>Sleep stage classification using single-channel EOG.</article-title> <source><italic>Comput. Biol. Med.</italic></source> <volume>102</volume> <fpage>211</fpage>&#x2013;<lpage>220</lpage>. <pub-id pub-id-type="doi">10.1016/j.compbiomed.2018.08.022</pub-id> <pub-id pub-id-type="pmid">30170769</pub-id></citation></ref>
<ref id="B61"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rechtschaffen</surname> <given-names>A.</given-names></name> <name><surname>Kales</surname> <given-names>A.</given-names></name></person-group> (<year>1968</year>). <source><italic>A manual of standardized terminology, techniques and scoring system for sleep stages of human subjects.</italic></source> <publisher-loc>Washington, DC</publisher-loc>: <publisher-name>US Government Printing Office</publisher-name>.</citation></ref>
<ref id="B62"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ribeiro</surname> <given-names>M.</given-names></name> <name><surname>Singh</surname> <given-names>S.</given-names></name> <name><surname>Guestrin</surname> <given-names>C.</given-names></name></person-group> (<year>2016</year>). &#x201C;<article-title>&#x201C;Why should i trust you?&#x201D; Explaining the predictions of any classifier</article-title>,&#x201D; in <source><italic>Proceedings of the ACM SIGKDD international conference on knowledge discovery and data mining</italic></source>, <publisher-loc>San Francisco, CA</publisher-loc>, <fpage>1135</fpage>&#x2013;<lpage>1144</lpage>. <pub-id pub-id-type="doi">10.1145/2939672.2939778</pub-id></citation></ref>
<ref id="B63"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rojas</surname> <given-names>I.</given-names></name> <name><surname>Joya</surname> <given-names>G.</given-names></name> <name><surname>Catala</surname> <given-names>A.</given-names></name></person-group> (<year>2017</year>). &#x201C;<article-title>Deep learning using EEG data in time and frequency domains for sleep stage classification</article-title>,&#x201D; in <source><italic>Proceedings of the IWANN 2017 advances in computational intelligence</italic></source>, <role>eds</role> <person-group person-group-type="editor"><name><surname>Rojas</surname> <given-names>I.</given-names></name> <name><surname>Joya</surname> <given-names>G.</given-names></name> <name><surname>Catala</surname> <given-names>A.</given-names></name></person-group> (<publisher-loc>Cham</publisher-loc>: <publisher-name>Springer</publisher-name>).</citation></ref>
<ref id="B64"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ruffini</surname> <given-names>G.</given-names></name> <name><surname>Iba&#x00F1;ez</surname> <given-names>D.</given-names></name> <name><surname>Castellano</surname> <given-names>M.</given-names></name> <name><surname>Dubreuil-Vall</surname> <given-names>L.</given-names></name> <name><surname>Soria-Frisch</surname> <given-names>A.</given-names></name> <name><surname>Postuma</surname> <given-names>R.</given-names></name><etal/></person-group> (<year>2019</year>). <article-title>Deep learning with EEG spectrograms in rapid eye movement behavior disorder.</article-title> <source><italic>Front. Neurol.</italic></source> <volume>10</volume>:<issue>806</issue>. <pub-id pub-id-type="doi">10.3389/fneur.2019.00806</pub-id> <pub-id pub-id-type="pmid">31417485</pub-id></citation></ref>
<ref id="B65"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Samek</surname> <given-names>W.</given-names></name> <name><surname>Binder</surname> <given-names>A.</given-names></name> <name><surname>Montavon</surname> <given-names>G.</given-names></name> <name><surname>Lapuschkin</surname> <given-names>S.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>K.</given-names></name></person-group> (<year>2017a</year>). <article-title>Evaluating the visualization of what a deep neural network has learned.</article-title> <source><italic>IEEE Trans. Neural. Netw. Learn. Syst.</italic></source> <volume>28</volume> <fpage>2660</fpage>&#x2013;<lpage>2673</lpage>. <pub-id pub-id-type="doi">10.1109/TNNLS.2016.2599820</pub-id> <pub-id pub-id-type="pmid">27576267</pub-id></citation></ref>
<ref id="B66"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Samek</surname> <given-names>W.</given-names></name> <name><surname>Wiegand</surname> <given-names>T.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>K.</given-names></name></person-group> (<year>2017b</year>). <article-title>Explainable artificial intelligence: Understanding, visualizing and interpreting deep learning models.</article-title> <source><italic>arXiv</italic></source> [<comment>Preprint</comment>]. <comment>arXiv:1708.08296</comment>.</citation></ref>
<ref id="B67"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Selvaraju</surname> <given-names>R.</given-names></name> <name><surname>Cogswell</surname> <given-names>M.</given-names></name> <name><surname>Das</surname> <given-names>A.</given-names></name> <name><surname>Vedantam</surname> <given-names>R.</given-names></name> <name><surname>Parikh</surname> <given-names>D.</given-names></name> <name><surname>Batra</surname> <given-names>D.</given-names></name></person-group> (<year>2020</year>). <article-title>Grad-CAM: Visual explanations from deep networks via gradient-based localization.</article-title> <source><italic>Int. J. Comput. Vis.</italic></source> <volume>128</volume> <fpage>336</fpage>&#x2013;<lpage>359</lpage>. <pub-id pub-id-type="doi">10.1007/s11263-019-01228-7</pub-id></citation></ref>
<ref id="B68"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Simonyan</surname> <given-names>K.</given-names></name> <name><surname>Vedaldi</surname> <given-names>A.</given-names></name> <name><surname>Zisserman</surname> <given-names>A.</given-names></name></person-group> (<year>2013</year>). <source><italic>Deep inside convolutional networks: Visualizing image classification models and saliency maps.</italic></source> Available online at: <ext-link ext-link-type="uri" xlink:href="http://arxiv.org/abs/1312.6034">http://arxiv.org/abs/1312.6034</ext-link> <comment>(accessed December 13, 2022)</comment>.</citation></ref>
<ref id="B69"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sors</surname> <given-names>A.</given-names></name> <name><surname>Bonnet</surname> <given-names>S.</given-names></name> <name><surname>Mirek</surname> <given-names>S.</given-names></name> <name><surname>Vercueil</surname> <given-names>L.</given-names></name> <name><surname>Payen</surname> <given-names>J.</given-names></name></person-group> (<year>2018</year>). <article-title>A convolutional neural network for sleep stage scoring from raw single-channel EEG.</article-title> <source><italic>Biomed. Signal. Process. Control</italic></source> <volume>42</volume> <fpage>107</fpage>&#x2013;<lpage>114</lpage>. <pub-id pub-id-type="doi">10.1016/j.bspc.2017.12.001</pub-id></citation></ref>
<ref id="B70"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sturm</surname> <given-names>I.</given-names></name> <name><surname>Lapuschkin</surname> <given-names>S.</given-names></name> <name><surname>Samek</surname> <given-names>W.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>K.</given-names></name></person-group> (<year>2016</year>). <article-title>Interpretable deep neural networks for single-trial EEG classification.</article-title> <source><italic>J. Neurosci. Methods</italic></source> <volume>274</volume> <fpage>141</fpage>&#x2013;<lpage>145</lpage>. <pub-id pub-id-type="doi">10.1016/j.jneumeth.2016.10.008</pub-id> <pub-id pub-id-type="pmid">27746229</pub-id></citation></ref>
<ref id="B71"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sullivan</surname> <given-names>H.</given-names></name> <name><surname>Schweikart</surname> <given-names>S.</given-names></name></person-group> (<year>2019</year>). <article-title>Are current tort liability doctrines adequate for addressing injury caused by AI?</article-title> <source><italic>AMA J. Ethics</italic></source> <volume>21</volume> <fpage>160</fpage>&#x2013;<lpage>166</lpage>. <pub-id pub-id-type="doi">10.1001/amajethics.2019.160</pub-id> <pub-id pub-id-type="pmid">30794126</pub-id></citation></ref>
<ref id="B72"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Supratak</surname> <given-names>A.</given-names></name> <name><surname>Dong</surname> <given-names>H.</given-names></name> <name><surname>Wu</surname> <given-names>C.</given-names></name> <name><surname>Guo</surname> <given-names>Y.</given-names></name></person-group> (<year>2017</year>). <article-title>DeepSleepNet: A model for automatic sleep stage scoring based on raw single-channel EEG.</article-title> <source><italic>IEEE Trans. Neural. Syst. Rehabil. Eng.</italic></source> <volume>25</volume> <fpage>1998</fpage>&#x2013;<lpage>2008</lpage>. <pub-id pub-id-type="doi">10.1109/TNSRE.2017.2721116</pub-id> <pub-id pub-id-type="pmid">28678710</pub-id></citation></ref>
<ref id="B73"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Thomas</surname> <given-names>A.</given-names></name> <name><surname>Heekeren</surname> <given-names>H.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>K.</given-names></name> <name><surname>Samek</surname> <given-names>W.</given-names></name></person-group> (<year>2018</year>). <source><italic>Analyzing neuroimaging data through recurrent deep learning models.</italic></source> Available online at: <ext-link ext-link-type="uri" xlink:href="http://arxiv.org/abs/1810.09945">http://arxiv.org/abs/1810.09945</ext-link> <comment>(accessed October 23, 2018)</comment>.</citation></ref>
<ref id="B74"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Thoret</surname> <given-names>E.</given-names></name> <name><surname>Andrillon</surname> <given-names>T.</given-names></name> <name><surname>Gauriau</surname> <given-names>C.</given-names></name> <name><surname>L&#x00E9;ger</surname> <given-names>D.</given-names></name> <name><surname>Pressnitzer</surname> <given-names>D.</given-names></name></person-group> (<year>2022</year>). <article-title>Sleep deprivation measured by voice analysis.</article-title> <source><italic>bioRxiv</italic></source> [<comment>Preprint</comment>]. <pub-id pub-id-type="doi">10.1101/2022.11.17.516913</pub-id></citation></ref>
<ref id="B75"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Thoret</surname> <given-names>E.</given-names></name> <name><surname>Andrillon</surname> <given-names>T.</given-names></name> <name><surname>L&#x00E9;ger</surname> <given-names>D.</given-names></name> <name><surname>Pressnitzer</surname> <given-names>D.</given-names></name></person-group> (<year>2021</year>). <article-title>Probing machine-learning classifiers using noise, bubbles, and reverse correlation.</article-title> <source><italic>J Neurosci Methods.</italic></source> <volume>362</volume>:<issue>109297</issue>. <pub-id pub-id-type="doi">10.1016/j.jneumeth.2021.109297</pub-id> <pub-id pub-id-type="pmid">34320410</pub-id></citation></ref>
<ref id="B76"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tsinalis</surname> <given-names>O.</given-names></name> <name><surname>Matthews</surname> <given-names>P.</given-names></name> <name><surname>Guo</surname> <given-names>Y.</given-names></name></person-group> (<year>2016a</year>). <article-title>Automatic sleep stage scoring using time-frequency analysis and stacked sparse autoencoders.</article-title> <source><italic>Ann. Biomed. Eng.</italic></source> <volume>44</volume> <fpage>1587</fpage>&#x2013;<lpage>1597</lpage>. <pub-id pub-id-type="doi">10.1007/s10439-015-1444-y</pub-id> <pub-id pub-id-type="pmid">26464268</pub-id></citation></ref>
<ref id="B77"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tsinalis</surname> <given-names>O.</given-names></name> <name><surname>Matthews</surname> <given-names>P.</given-names></name> <name><surname>Guo</surname> <given-names>Y.</given-names></name> <name><surname>Zafeiriou</surname> <given-names>S.</given-names></name></person-group> (<year>2016b</year>). <article-title>Automatic sleep stage scoring with single-channel EEG using convolutional neural networks.</article-title> <source><italic>arXiv</italic></source> [<comment>Preprint</comment>]. Available from: <ext-link ext-link-type="uri" xlink:href="http://arxiv.org/abs/1610.01683">http://arxiv.org/abs/1610.01683</ext-link> <comment>(accessed December 13, 2022)</comment>.</citation></ref>
<ref id="B78"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tuk</surname> <given-names>B.</given-names></name> <name><surname>Obery&#x00E9;</surname> <given-names>J.</given-names></name> <name><surname>Pieters</surname> <given-names>M.</given-names></name> <name><surname>Schoemaker</surname> <given-names>R.</given-names></name> <name><surname>Kemp</surname> <given-names>B.</given-names></name> <name><surname>Van Gerven</surname> <given-names>J.</given-names></name><etal/></person-group> (<year>1997</year>). <article-title>Pharmacodynamics of temazepam in primary insomnia: Assessment of the value of quantitative electrocephalography and saccadic eye movements in predicting improvement of sleep.</article-title> <source><italic>Clin. Pharmacol. Ther.</italic></source> <volume>62</volume> <fpage>444</fpage>&#x2013;<lpage>452</lpage>. <pub-id pub-id-type="doi">10.1016/S0009-9236(97)90123-5</pub-id> <pub-id pub-id-type="pmid">9357396</pub-id></citation></ref>
<ref id="B79"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Van Sweden</surname> <given-names>B.</given-names></name> <name><surname>Kemp</surname> <given-names>B.</given-names></name> <name><surname>Kamphuisen</surname> <given-names>H.</given-names></name> <name><surname>Van der Velde</surname> <given-names>E.</given-names></name></person-group> (<year>1990</year>). <article-title>Alternative electrode placement in (automatic) sleep scoring (F(pz)-C(z)/P(z)-O(z) versus C(4)-A(1)).</article-title> <source><italic>Sleep</italic></source> <volume>13</volume> <fpage>279</fpage>&#x2013;<lpage>283</lpage>. <pub-id pub-id-type="doi">10.1093/sleep/13.3.279</pub-id> <pub-id pub-id-type="pmid">2356401</pub-id></citation></ref>
<ref id="B80"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Vilamala</surname> <given-names>A.</given-names></name> <name><surname>Madsen</surname> <given-names>K.</given-names></name> <name><surname>Hansen</surname> <given-names>L.</given-names></name></person-group> (<year>2017</year>). &#x201C;<article-title>Deep convolutional neural networks for interpretable analysis of EEG sleep stage scoring</article-title>,&#x201D; in <source><italic>Proceedings of the IEEE 2017 international workshop on machine learning for signal processing</italic></source>, <publisher-loc>Tokyo</publisher-loc>. <pub-id pub-id-type="doi">10.1109/MLSP.2017.8168133</pub-id></citation></ref>
<ref id="B81"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>I.</given-names></name> <name><surname>Lee</surname> <given-names>C.</given-names></name> <name><surname>Kim</surname> <given-names>H.</given-names></name> <name><surname>Kim</surname> <given-names>H.</given-names></name> <name><surname>Kim</surname> <given-names>D.</given-names></name></person-group> (<year>2020</year>). &#x201C;<article-title>An ensemble deep learning approach for sleep stage classification via single-channel EEG and EOG</article-title>,&#x201D; in <source><italic>Proceedings of the 11th international conference on ICT convergence: Data, network, and AI in the age of Untact</italic></source>, <publisher-loc>Washington, DC</publisher-loc>, <fpage>394</fpage>&#x2013;<lpage>398</lpage>. <pub-id pub-id-type="doi">10.1109/ICTC49870.2020.9289335</pub-id></citation></ref>
<ref id="B82"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yan</surname> <given-names>W.</given-names></name> <name><surname>Plis</surname> <given-names>S.</given-names></name> <name><surname>Calhoun</surname> <given-names>V.</given-names></name> <name><surname>Liu</surname> <given-names>S.</given-names></name> <name><surname>Jiang</surname> <given-names>R.</given-names></name> <name><surname>Jiang</surname> <given-names>T.</given-names></name><etal/></person-group> (<year>2017</year>). &#x201C;<article-title>Discriminating schizophrenia from normal controls using resting state functional network connectivity: A deep neural network and layer-wise relevance propagation method</article-title>,&#x201D; in <source><italic>Proceedings of the IEEE international workshop on machine learning for signal processing</italic></source>, <publisher-loc>Tokyo</publisher-loc>. <pub-id pub-id-type="doi">10.1109/MLSP.2017.8168179</pub-id></citation></ref>
<ref id="B83"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Youness</surname> <given-names>M.</given-names></name></person-group> (<year>2020</year>). <source><italic>CVxTz/EEG\_classification: v1.0.</italic></source> Available from: <ext-link ext-link-type="uri" xlink:href="https://github.com/CVxTz/EEG_classification">https://github.com/CVxTz/EEG_classification</ext-link> <comment>(accessed January 5, 2021)</comment>.</citation></ref>
<ref id="B84"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhai</surname> <given-names>B.</given-names></name> <name><surname>Perez-Pozuelo</surname> <given-names>I.</given-names></name> <name><surname>Clifton</surname> <given-names>E.</given-names></name> <name><surname>Palotti</surname> <given-names>J.</given-names></name> <name><surname>Guan</surname> <given-names>Y.</given-names></name></person-group> (<year>2020</year>). <article-title>Making sense of sleep: Multimodal sleep stage classification in a large, diverse population using movement and cardiac sensing.</article-title> <source><italic>Proc. ACM Interact. Mobile Wearable Ubiquitous Technol.</italic></source> <volume>4</volume> <fpage>1</fpage>&#x2013;<lpage>33</lpage>. <pub-id pub-id-type="doi">10.1145/3397325</pub-id></citation></ref>
<ref id="B85"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>D.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Zhou</surname> <given-names>L.</given-names></name> <name><surname>Yuan</surname> <given-names>H.</given-names></name> <name><surname>Shen</surname> <given-names>D.</given-names></name></person-group> (<year>2011</year>). <article-title>Multimodal classification of Alzheimer&#x2019;s disease and mild cognitive impairment.</article-title> <source><italic>Neuroimage</italic></source> <volume>55</volume> <fpage>856</fpage>&#x2013;<lpage>867</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuroimage.2011.01.008</pub-id> <pub-id pub-id-type="pmid">21236349</pub-id></citation></ref>
</ref-list>
</back>
</article>
