<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Neurosci.</journal-id>
<journal-title>Frontiers in Neuroscience</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Neurosci.</abbrev-journal-title>
<issn pub-type="epub">1662-453X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fnins.2024.1368641</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Neuroscience</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Effect of spectral degradation on speech intelligibility and cortical representation</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Choi</surname> <given-names>Hyo Jung</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="author-notes" rid="fn0001"><sup>&#x2020;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2102524/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Kyong</surname> <given-names>Jeong-Sug</given-names></name>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<xref ref-type="author-notes" rid="fn0001"><sup>&#x2020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Won</surname> <given-names>Jong Ho</given-names></name>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/215655/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Shim</surname> <given-names>Hyun Joon</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/675985/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Department of Otorhinolaryngology-Head and Neck Surgery, Nowon Eulji Medical Center, Eulji University School of Medicine</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country></aff>
<aff id="aff2"><sup>2</sup><institution>Eulji Tinnitus and Hearing Research Institute, Nowon Eulji Medical Center</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country></aff>
<aff id="aff3"><sup>3</sup><institution>Sensory-Organ Research Institute, Medical Research Center, Seoul National University School of Medicine</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country></aff>
<aff id="aff4"><sup>4</sup><institution>Department of Radiology, Konkuk University Medical Center</institution>, <addr-line>Seoul</addr-line>, <country>Republic of Korea</country></aff>
<aff id="aff5"><sup>5</sup><institution>Hyman, Phelps and McNamara, P.C.</institution>, <addr-line>Washington, DC</addr-line>, <country>United States</country></aff>
<author-notes>
<fn fn-type="edited-by" id="fn0002">
<p>Edited by: Mathieu Bourguignon, Universit&#x00E9; libre de Bruxelles, Belgium</p>
</fn>
<fn fn-type="edited-by" id="fn0003">
<p>Reviewed by: I-Hui Hsieh, National Central University, Taiwan</p>
<p>Ting-Ting Chang, National Chengchi University, Taiwan</p>
</fn>
<corresp id="c001">&#x002A;Correspondence: Hyun Joon Shim, <email>eardoc11@naver.com</email></corresp>
<fn fn-type="equal" id="fn0001">
<p><sup>&#x2020;</sup>These authors share first authorship</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>05</day>
<month>04</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>18</volume>
<elocation-id>1368641</elocation-id>
<history>
<date date-type="received">
<day>10</day>
<month>01</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>25</day>
<month>03</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2024 Choi, Kyong, Won and Shim.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Choi, Kyong, Won and Shim</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Noise-vocoded speech has long been used to investigate how acoustic cues affect speech understanding. Studies indicate that reducing the number of spectral channel bands diminishes speech intelligibility. Despite previous studies examining the channel band effect using earlier event-related potential (ERP) components, such as P1, N1, and P2, a clear consensus or understanding remains elusive. Given our hypothesis that spectral degradation affects higher-order processing of speech understanding beyond mere perception, we aimed to objectively measure differences in higher-order abilities to discriminate or interpret meaning. Using an oddball paradigm with speech stimuli, we examined how neural signals correlate with the evaluation of speech stimuli based on the number of channel bands measuring N2 and P3b components. In 20 young participants with normal hearing, we measured speech intelligibility and N2 and P3b responses using a one-syllable task paradigm with animal and non-animal stimuli across four vocoder conditions with 4, 8, 16, or 32 channel bands. Behavioral data from word repetition clearly affected the number of channel bands, and all pairs were significantly different (<italic>p</italic>&#x2009;&#x003C;&#x2009;0.001). We also observed significant effects of the number of channels on the peak amplitude [<italic>F</italic><sub>(2.006, 38.117)</sub>&#x2009;=&#x2009;9.077, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001] and peak latency [<italic>F</italic><sub>(3, 57)</sub>&#x2009;=&#x2009;26.642, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001] of the N2 component. Similarly, the P3b component showed significant main effects of the number of channel bands on the peak amplitude [<italic>F</italic><sub>(2.231, 42.391)</sub>&#x2009;=&#x2009;13.045, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001] and peak latency [<italic>F</italic><sub>(3, 57)</sub>&#x2009;=&#x2009;2.968, <italic>p</italic>&#x2009;=&#x2009;0.039]. In summary, our findings provide compelling evidence that spectral channel bands profoundly influence cortical speech processing, as reflected in the N2 and P3b components, a higher-order cognitive process. We conclude that spectrally degraded one-syllable speech primarily affects cortical responses during semantic integration.</p>
</abstract>
<kwd-group>
<kwd>speech intelligibility</kwd>
<kwd>spectral degradation</kwd>
<kwd>vocoder</kwd>
<kwd>event-related potential</kwd>
<kwd>N2 and P3b</kwd>
</kwd-group>
<counts>
<fig-count count="6"/>
<table-count count="4"/>
<equation-count count="0"/>
<ref-count count="42"/>
<page-count count="11"/>
<word-count count="6737"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Auditory Cognitive Neuroscience</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec1">
<label>1</label>
<title>Introduction</title>
<p>Spectral degradation of speech limits successful comprehension. Informational or energetic noise can mask speech by degrading spectral information and thus reduce its intelligibility. Noise-vocoded speech is a form of spectrally degraded speech with distortion that was developed by Shannon et al. to simulate speech heard through a cochlear implant (CI) device (<xref ref-type="bibr" rid="ref34">Shannon et al., 1995</xref>). The basic principle of vocoder processing is to decompose a speech signal into multiple bands, extracting the envelope waveform for each band modulating a carrier signal and then summate all amplitude-modulated carrier signals (<xref ref-type="bibr" rid="ref34">Shannon et al., 1995</xref>; <xref ref-type="bibr" rid="ref11">Dorman et al., 1997</xref>). Until now, CIs have comprised at most 40 channels (Advanced Bionics recently released implants with 120 spectral bands), and the technical limitations prevent a literal copy of the human cochlear, which can detect thousands of tonotopy resolutions across the range of 20&#x2013;20,000&#x2009;Hz. Accordingly, CI users suffer from hearing spectrally degraded speech.</p>
<p>Numerous studies that measured speech intelligibility consistently observed a clear stepwise increase in the number of correctly reported words with increasing number of channel bands in noise-vocoded speech (<xref ref-type="bibr" rid="ref18">Hervais-Adelman et al., 2008</xref>; <xref ref-type="bibr" rid="ref35">Souza and Rosen, 2009</xref>). Along with these behavioral data, the effect of number of channel bands on spectral information was also observed in several functional MRI studies. Increased intelligibility led to an increased percent change in the blood oxygen level-dependent signal in the temporal lobe (<xref ref-type="bibr" rid="ref6">Davis and Johnsrude, 2003</xref>; <xref ref-type="bibr" rid="ref28">Obleser et al., 2007</xref>; <xref ref-type="bibr" rid="ref12">Evans et al., 2014</xref>).</p>
<p>While several electroencephalography (EEG) studies have replicated the influence of channel bands on cortical potentials, the current body of evidence is insufficient to draw definitive conclusions. Despite previous studies examining the channel band effect using earlier event-related potentials (ERP) components, such as P1, N1, and P2, a clear consensus or understanding remains elusive. In the study of <xref ref-type="bibr" rid="ref16">Friesen et al. (2009)</xref>, the effects of channel bands were reflected in the P1, N1, and P2 components in conditions involving 2, 4, 8, 12, and 16 channels, as well as ordinary speech. The study revealed that, as the number of channels of acoustic information increased, the peak amplitude of the neural response increased. In contrast, a recent study found inconsistent changes in the P1, N1, and P2 components when listening to vocoded speech on 4 and 22 channels (<xref ref-type="bibr" rid="ref10">Dong and Gai, 2021</xref>). In another study, the type of vocoder carrier was shown to affect the neural responses; the responses to noise-vocoded stimuli were smaller and slower when measured using mismatch negativity compared with the responses to tone-vocoded stimuli (<xref ref-type="bibr" rid="ref41">Xu et al., 2019</xref>). The P1, N1, and P2 components assessed in those studies focused on the neural responses derived from the physical or acoustic perception of speech. However, listening to speech with sparse information increases the compensatory reliance on top-down cognitive processes (<xref ref-type="bibr" rid="ref30">Pals et al., 2020</xref>). A contribution of the frontal lobe was also evident, and the frontal operculum showed an elevated response to noise-vocoded speech (<xref ref-type="bibr" rid="ref6">Davis and Johnsrude, 2003</xref>). Inconsistent findings and limitations in capturing higher-level cognitive processes for the earlier ERPs prompted a shift in focus to the N2 and P3 components in our study, rather than the P1, N1, and P2 components, as cortical potentials to replicate the channel band effect. We hypothesized that the N2 and P3 components, associated with lexical information assessment and stimulus categorization, would offer a more detailed exploration of the channel band effect in noise-vocoded speech.</p>
<p>N2 and P3 deflections are elicited in response to task-relevant stimuli in an oddball paradigm (<xref ref-type="bibr" rid="ref24">Luck, 2014</xref>). The N2 component refers to a frontocentral negativity occurring 200 to 350&#x2009;ms after stimuli (<xref ref-type="bibr" rid="ref33">Schmitt et al., 2000</xref>; <xref ref-type="bibr" rid="ref15">Folstein and Van Petten, 2008</xref>). The latency, however, is typically delayed in tasks with complex stimuli, such as speech words, encompassing a broader time window (350 to 800&#x2009;ms), depending on task difficulty or hearing condition. Thus, this prolonged N2 is often referred to as N2N4 (<xref ref-type="bibr" rid="ref40">Voola et al., 2023</xref>), reflecting cortical access to lexical information and semantic categorization in the deaf population (<xref ref-type="bibr" rid="ref14">Finke et al., 2016</xref>). Both N2 and N4 are cortical responses related to the lexical selection process, analogous to the functional interpretation of the N400 component (<xref ref-type="bibr" rid="ref37">Van den Brink and Hagoort, 2004</xref>). P3 has subcomponents of P3a and P3b. Unlike P3a, P3b reflects the effortful allocation of resources to discriminate and interpret auditory stimuli (<xref ref-type="bibr" rid="ref38">Volpe et al., 2007</xref>; <xref ref-type="bibr" rid="ref39">Voola et al., 2022</xref>). The P3b component is also associated with updating working memory, and prolonged latencies may be interpreted as slower stimulus evaluation (<xref ref-type="bibr" rid="ref5">Beynon et al., 2005</xref>; <xref ref-type="bibr" rid="ref17">Henkin et al., 2015</xref>). The P3b component is maximally observed parietally, elicited by unpredictable and infrequent shifts. Any manipulation to delay stimulus categorization increases P3b latency and decreases amplitude (<xref ref-type="bibr" rid="ref21">Johnson, 1988</xref>). Studies have shown that both N2 (<xref ref-type="bibr" rid="ref14">Finke et al., 2016</xref>) and P3b (<xref ref-type="bibr" rid="ref5">Beynon et al., 2005</xref>) were prolonged in CI users compared to normal hearing listeners, implicating a slower stimulus evaluation in CI users due to device limitations such as poor spectral density. However, there is no report on the N2 and P3b regarding the effect of spectral manipulation in speech.</p>
<p>In the current study, we aimed to investigate differences in acoustic-based semantic processing in noise-vocoded speech. We employed a one-syllable oddball paradigm instead of sentences, which allowed us to: (1) minimize the redundancy of cues, (2) reduce top-down expectations in the context (<xref ref-type="bibr" rid="ref3">Bae et al., 2022</xref>), and (3) control for individual differences in education and attention ability. Stimulus categorization, such as determining whether a stimulus is living or non-living, represents one of the simplest forms of higher-order processing in speech. We hypothesized that measuring the N2 and P3b components through the vocoded one-syllable oddball paradigm could robustly evaluate the impact of spectral degradation on acoustic-based semantic processing. The purpose of this study was to objectively assess the effect of channel bands, focusing on higher-level cognition such as lexical information assessment and stimulus categorization.</p>
</sec>
<sec id="sec2">
<label>2</label>
<title>Subjects and methods</title>
<sec id="sec3">
<label>2.1</label>
<title>Subjects</title>
<p>The main experiment&#x2019;s sample size was determined through a pilot study with four participants. The effect size (<italic>&#x03B7;p<sup>2</sup></italic>) was obtained through a pilot study and using G&#x002A;Power software (latest version 3.1.9.7; Heinrich-Heine-Universit&#x00E4;t D&#x00FC;sseldorf, D&#x00FC;sseldorf, Germany; <xref ref-type="bibr" rid="ref1001">Faul et al., 2007</xref>), the recommended sample size of 17 was derived by inputting the effect size into the software. To account for potential dropouts (20%), we aimed to recruit 21 participants. Despite one withdrawal, data analysis was ultimately conducted with a total of 20 young adults with normal hearing (mean age: 29.8&#x2009;&#x00B1;&#x2009;5.9&#x2009;years old; women, 29.0&#x2009;&#x00B1;&#x2009;6.5&#x2009;years old; men, 30.6&#x2009;&#x00B1;&#x2009;5.6&#x2009;years; 10 males, 10 females).</p>
<p>The pure-tone average across 500, 1,000, 2,000, and 3,000&#x2009;Hz was 7.7 [standard deviation (SD) =3.0] decibel (dB) hearing level (HL) on the right side and 6.8 (SD&#x2009;=&#x2009;2.9) dB HL on the left side. <xref ref-type="table" rid="tab1">Table 1</xref> shows the detailed demographic information of the participants. The study was conducted in accordance with the Declaration of Helsinki and the recommendations of the Institutional Review Board of &#x002A;&#x002A;&#x002A;&#x002A; Medical Center. Written informed consent was obtained from all subjects. After subjects signed the consent form, a copy was given to them.</p>
<table-wrap position="float" id="tab1">
<label>Table 1</label>
<caption>
<p>Demographic summary of the participants.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">No.</th>
<th align="left" valign="top">Sex</th>
<th align="center" valign="top">Age (years)</th>
<th align="center" valign="top">Handedness</th>
<th align="center" valign="top">Education (years)</th>
<th align="center" valign="top">Pure tone average (right, dB HL)</th>
<th align="center" valign="top">Pure tone average (left, dB HL)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">1</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">38</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">18</td>
<td align="center" valign="middle">8.8</td>
<td align="center" valign="middle">8.8</td>
</tr>
<tr>
<td align="left" valign="middle">2</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">34</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">12</td>
<td align="center" valign="middle">10.0</td>
<td align="center" valign="middle">10.0</td>
</tr>
<tr>
<td align="left" valign="middle">3</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">36</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">8.8</td>
<td align="center" valign="middle">10.0</td>
</tr>
<tr>
<td align="left" valign="middle">4</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">24</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">8.8</td>
<td align="center" valign="middle">6.3</td>
</tr>
<tr>
<td align="left" valign="middle">5</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">31</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">7.5</td>
<td align="center" valign="middle">8.8</td>
</tr>
<tr>
<td align="left" valign="middle">6</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">20</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">12</td>
<td align="center" valign="middle">5.0</td>
<td align="center" valign="middle">5.0</td>
</tr>
<tr>
<td align="left" valign="middle">7</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">21</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">12</td>
<td align="center" valign="middle">5.0</td>
<td align="center" valign="middle">5.0</td>
</tr>
<tr>
<td align="left" valign="middle">8</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">36</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">5.0</td>
<td align="center" valign="middle">5.0</td>
</tr>
<tr>
<td align="left" valign="middle">9</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">27</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">14</td>
<td align="center" valign="middle">6.3</td>
<td align="center" valign="middle">8.8</td>
</tr>
<tr>
<td align="left" valign="middle">10</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">26</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">14</td>
<td align="center" valign="middle">8.8</td>
<td align="center" valign="middle">3.8</td>
</tr>
<tr>
<td align="left" valign="middle">11</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">23</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">12</td>
<td align="center" valign="middle">11.3</td>
<td align="center" valign="middle">11.3</td>
</tr>
<tr>
<td align="left" valign="middle">12</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">36</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">18</td>
<td align="center" valign="middle">2.5</td>
<td align="center" valign="middle">6.3</td>
</tr>
<tr>
<td align="left" valign="middle">13</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">27</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">18</td>
<td align="center" valign="middle">3.8</td>
<td align="center" valign="middle">2.5</td>
</tr>
<tr>
<td align="left" valign="middle">14</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">32</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">10.0</td>
<td align="center" valign="middle">7.5</td>
</tr>
<tr>
<td align="left" valign="middle">15</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">29</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">11.3</td>
<td align="center" valign="middle">7.5</td>
</tr>
<tr>
<td align="left" valign="middle">16</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">23</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">6.3</td>
<td align="center" valign="middle">0.0</td>
</tr>
<tr>
<td align="left" valign="middle">17</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">39</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">14</td>
<td align="center" valign="middle">5.0</td>
<td align="center" valign="middle">5.0</td>
</tr>
<tr>
<td align="left" valign="middle">18</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">32</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">15.0</td>
<td align="center" valign="middle">11.3</td>
</tr>
<tr>
<td align="left" valign="middle">19</td>
<td align="left" valign="bottom">M</td>
<td align="center" valign="bottom">27</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">16</td>
<td align="center" valign="middle">7.5</td>
<td align="center" valign="middle">7.5</td>
</tr>
<tr>
<td align="left" valign="middle">20</td>
<td align="left" valign="bottom">F</td>
<td align="center" valign="bottom">35</td>
<td align="center" valign="bottom">R</td>
<td align="center" valign="middle">18</td>
<td align="center" valign="middle">7.5</td>
<td align="center" valign="middle">5.0</td>
</tr>
<tr>
<td align="left" valign="middle">Mean</td>
<td/>
<td align="center" valign="bottom">29.8</td>
<td/>
<td align="center" valign="middle">15.3</td>
<td align="center" valign="middle">7.7</td>
<td align="center" valign="middle">6.8</td>
</tr>
<tr>
<td align="left" valign="middle">SD</td>
<td/>
<td align="center" valign="bottom">5.9</td>
<td/>
<td align="center" valign="middle">2.1</td>
<td align="center" valign="middle">3.0</td>
<td align="center" valign="middle">2.9</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>The pure-tone average was calculated as the average of the thresholds at frequencies of 500, 1,000, 2,000, and 3,000&#x2009;Hz.</p>
</table-wrap-foot>
</table-wrap>
</sec>
<sec id="sec4">
<label>2.2</label>
<title>Stimuli</title>
<p>Stimuli were recorded by a female speaker reading five lists of 25 monosyllable animal or non-animal Korean words using a lapel microphone (BY-WMA4 PRO K3, BOYA, Shenzhen, Hong Kong) in a soundproof booth. All the recorded stimuli were sampled at a rate of 44,100&#x2009;Hz, and the overall root mean square amplitude was set at &#x2212;22&#x2009;dB. The phonetic balance, an equal range of the phonetic composition of speech, words in common usage, and familiarity with the words were considered when the word lists were chosen. Based on the frequency of occurrence of conversational sounds, we created a CVC word set with 19 initial consonants, 21 vowels, and 7 final consonants.</p>
<p>The equivalent average difficulty and phoneme composition of the lists were verified. The long-term average speech spectrum of the recorded syllable was analyzed using Computerized Speech Lab (CSL model 4500b, KayPENTAX Elemetrics Corporation, Lincoln Park, NJ, United States).</p>
<p>Noise-vocoding involves passing a speech signal through a filter bank to extract time-varying envelopes associated with the energy in each spectral channel band. The extracted envelopes were multiplied by white noise and combined after re-filtering (<xref ref-type="bibr" rid="ref34">Shannon et al., 1995</xref>). First, the initial signal underwent processing through band-pass filtering, creating multiple channels (4, 8, 16, or 32 channels). The cut-off frequencies for each individual band-pass filter were determined using logarithmically spaced frequency bands, employing the Greenwood function (e.g., for 4 channels: [80, 424, 1,250, 3,234, and 8,000&#x2009;Hz]). The center frequency of each channel was computed as the geometric mean between the two cutoff frequencies associated with that specific channel. The collective input frequency range spanned from 80 to 8,000&#x2009;Hz. Subsequently, the amplitude envelope was extracted for each frequency band by means of half-wave rectification. Finally, we then summed the signals to generate the noise-vocoded speech session (<xref ref-type="bibr" rid="ref34">Shannon et al., 1995</xref>; <xref ref-type="bibr" rid="ref13">Faulkner et al., 2012</xref>; <xref ref-type="bibr" rid="ref12">Evans et al., 2014</xref>). Vocoding was performed using a custom MATLAB script (2020a, MathWorks, Inc., Natick, MA, United States) using 4, 8, 16, or 32 spectral channels with a temporal envelope modulation cut-off frequency fixed at 500&#x2009;Hz. <xref ref-type="fig" rid="fig1">Figure 1</xref> illustrates the flow of generating noise-vocoded signal. Noise-vocoded speech sounds like a harsh whisper with only a weak sense of pitch. From our lab experience, speech synthesized with fewer than 4 bands were hard to understand. Resulting signals sound like a harsh whisper (<xref ref-type="bibr" rid="ref26">Narain et al., 2003</xref>), and spectral detail decreases as the number of channel band decreases, as seen in <xref ref-type="fig" rid="fig2">Figure 2</xref>.</p>
<fig position="float" id="fig1">
<label>Figure 1</label>
<caption>
<p>Illustration depicting the generation of the noise-vocoded signal. The input signals were band-pass filtered into 4 (BPF1), 8 (BPF2), 16 (BPH3), and 32 (BPF4) channel bands prior to Hilbert transformation. After separating the envelopes from the temporal fine structures, the vocoder speech signal was generated by adding a noise carrier to the envelopes.</p>
</caption>
<graphic xlink:href="fnins-18-1368641-g001.tif"/>
</fig>
<fig position="float" id="fig2">
<label>Figure 2</label>
<caption>
<p>Spectrogram of the signals of the conditions with 4, 8, 16, and 32 channel bands. With fewer channel bands, the speech becomes more spectrally degraded and harder to understand. The information in the spectra was blurriest in the condition with 4 channel bands.</p>
</caption>
<graphic xlink:href="fnins-18-1368641-g002.tif"/>
</fig>
</sec>
<sec id="sec5">
<label>2.3</label>
<title>Procedure</title>
<sec id="sec6">
<label>2.3.1</label>
<title>Behavioral test</title>
<p>The perception of one-syllable words was tested in four different channel band conditions (4, 8, 16, or 32 channel vocoder), using five lists, with each containing 25 Korean monosyllabic words. The participants were asked to repeat the words after they were presented through a loudspeaker placed 1 meter in front of them. The stimulus intensity was set to 70&#x2009;dB sound pressure level (SPL) when calibrated at the listener&#x2019;s head position, 1 meter away from the loudspeaker. In addition, during the recording of the N2 and P3b, we assessed participants&#x2019; accuracy in identifying one-syllable non-animal words from animal/non-animal word sets that were vocoded in four different channel band conditions (4, 8, 16, and 32 channel bands).</p>
</sec>
<sec id="sec7">
<label>2.3.2</label>
<title>EEG</title>
<p>The neural responses were recorded across 31 AG-Ag/Cl sintered electrodes placed according to the international 10&#x2013;20 system (<xref ref-type="bibr" rid="ref20">Jasper, 1958</xref>) in an elastic cap using the actiCHamp Brain Products recording system (BrainVision Recorder Professional, V.1.23.0001, Brain Products GmbH, Inc., Munich, Germany) while the participant sat in a dimly lit, sound-attenuated, electrically soundproof booth. Electro-oculogram and electrocardiogram were also tagged to trace eye movements and heartbeats. EEG data were digitized online at a sampling rate of 1,000&#x2009;Hz. All 32 electrodes were referenced to the algebraic average of all electrodes/channels and were therefore unbiased to any electrode position. The ground electrode was placed between electrodes Fp1 and Fp2. Software filters were set at low (0.5&#x2009;Hz) and high (70&#x2009;Hz) cutoffs. A notch filter was set at 60&#x2009;Hz to prevent powerline noise. The impedance of each scalp electrode was kept below 5 k&#x03A9; throughout the recording, as suggested by the manufacturer&#x2019;s guide.</p>
<sec id="sec8">
<label>2.3.2.1</label>
<title>Oddball paradigm</title>
<p>Based on a one-syllable task paradigm, the participants listened to animal words or non-animal but sensible words. Overall, 70% of the trials involved an animal target word (e.g., mouse, snake, bear; all monosyllabic in Korean). Sitting upright, the listeners were instructed that monosyllabic animal words were heard but they were to push the button otherwise as soon as possible. This controlled for the participant&#x2019;s attentive level during recording, and this is for recording behavioral performance identifying non-animal words from animal words. The participants were randomly presented with six blocks of 210 animal words and 90 non-animal words in four channel band conditions, totaling 1,200 trials. The interstimulus interval was fixed in 2,000&#x2009;ms, allowing for a jitter of 2&#x2013;5&#x2009;ms. The order of presentation was randomized within each block, and the order of blocks was counterbalanced among listeners using E-Prime software (version 3, Psychology Software Tools, Inc., Sharpsburg, PA). Each block was separated by a 2~5-min break. A familiarizing session ensured the participants understood the task and their muscles were sufficiently relaxed. The intensity of the sound was fixed approximately at 70&#x2009;dB SPL when calibrated at the listener&#x2019;s head position, 1 meter from the loudspeaker.</p>
</sec>
<sec id="sec9">
<label>2.3.2.2</label>
<title>Preprocessing of the neural signals</title>
<p>The data were preprocessed and analyzed using Brain Vision analyzer (version 2.0, Brain Products GmbH, Inc.) and MATLAB R2019b (MathWorks Inc.) with EEGLAB v2021 (<xref ref-type="bibr" rid="ref9">Delorme and Makeig, 2004</xref>), and Fieldtrip (<xref ref-type="bibr" rid="ref29">Oostenveld et al., 2011</xref>) toolboxes. The EEG was filtered with a 0.1&#x2009;Hz high-pass filter (Butterworth with a 12&#x2009;dB/octave roll-off) and low-pass filtered at 50&#x2009;Hz (Butterworth with a 24&#x2009;dB/octave roll-off). The first three trials were excluded from analyses. The data were resampled at 256&#x2009;Hz. Independent component analysis (ICA) was used to reject artifacts associated with eye blinks and body movement (average 4 independent components, range 3&#x2013;6) and reconstructed (<xref ref-type="bibr" rid="ref25">Makeig et al., 1997</xref>; <xref ref-type="bibr" rid="ref22">Jung et al., 2000</xref>), transforming to the average reference. EEG waveforms were then time-locked to each stimulus onset and segmented from 200&#x2009;ms prior to the stimulus onset to 1,000&#x2009;ms after the stimulus onset. Baseline correction was performed accordingly. Prior to averaging, bad channels were interpolated using a spherical spline function (<xref ref-type="bibr" rid="ref32">Perrin et al., 1989</xref>), and segments with values greater than &#x00B1;70&#x2009;&#x03BC;V at any electrode were rejected. All participants had at least 180&#x2013;200 out of 210 usable animal trials and 78&#x2013;86 usable non-animal trials per vocoder channel-band condition. An average wave file was generated for each subject for each condition. Based on the grand average computed across all conditions and participants, latency ranges for N2 and P3b were determined according to the literature and the peak latency was measured using a half area quantification, which may be less affected by latency jitter (<xref ref-type="bibr" rid="ref24">Luck, 2014</xref>; <xref ref-type="bibr" rid="ref14">Finke et al., 2016</xref>). Difference waveforms were constructed based on the subtraction of target stimuli from standard stimuli within conditions (<xref ref-type="bibr" rid="ref8">Deacon et al., 1991</xref>). The area latency and amplitude of the N2 and P3b difference waveforms at each condition were compared. The time windows for N2 and P3b analysis were defined from each average waveform. In our data, the time windows for N2 and P3b were set as 280&#x2013;870&#x2009;ms and 280&#x2013;840&#x2009;ms, respectively. N2 was measured by averaging the signals from the frontocentral electrodes (Fz, FC1, FC2, and Cz), while P3b was measured using the parietal electrodes (CP1, CP2, P3, P4, and Pz), as outlined in <xref ref-type="bibr" rid="ref14">Finke et al. (2016)</xref>.</p>
</sec>
</sec>
</sec>
<sec id="sec10">
<label>2.4</label>
<title>Statistical analysis</title>
<p>Repeated-measures analysis of variance (RM ANOVA) was used to test the effects of vocoder channel-band on behavioral accuracy and N2 and P3b area peak amplitude and latency. Greenhouser&#x2013;Geisser correction was applied to the statistical comparisons for which the Mauchly test indicated violation of sphericity. Follow-up <italic>post hoc</italic> Bonferroni corrected tests were also used. Significance was inferred for corrected <italic>p</italic>-values of &#x003C;0.05. Correlation analyses were performed using Pearson&#x2019;s correlation test. All statistical analyses were performed using IBM SPSS software (ver. 25.0; IBM Corp, Armonk, NY, United States) and the built-in functions in MATLAB (2014a, 2019a, MathWorks Inc.). Data are presented as the mean&#x2009;&#x00B1;&#x2009;SD. Outliers were defined as values that differed from the mean by &#x00B1;2 SD.</p>
</sec>
</sec>
<sec sec-type="results" id="sec11">
<label>3</label>
<title>Results</title>
<sec id="sec12">
<label>3.1</label>
<title>Behavioral data</title>
<p>In vocoded speech perception, the accuracy was 3.20%&#x2009;&#x00B1;&#x2009;5.27, 35.80%&#x2009;&#x00B1;&#x2009;8.10, 52.00%&#x2009;&#x00B1;&#x2009;9.94, and 63.60%&#x2009;&#x00B1;&#x2009;9.88 in 4, 8, 16, and 32 channel band conditions, respectively. RM ANOVA showed a significant effect of the number of channel bands [<italic>F</italic><sub>(3, 57)</sub>&#x2009;=&#x2009;284.70, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001, <italic>&#x03B7;p<sup>2</sup></italic>&#x2009;=&#x2009;0.937]. All the pairs were significantly different (<italic>P<sub>Bonf</sub></italic>&#x2009;&#x003C;&#x2009;0.001; <xref ref-type="fig" rid="fig3">Figure 3</xref>; <xref ref-type="table" rid="tab2">Table 2</xref>).</p>
<fig position="float" id="fig3">
<label>Figure 3</label>
<caption>
<p>Vocoded speech perceptions. The conditions with fewer channel bands scored lower than conditions with more channel bands. All of the pairs were significantly different from each other. <sup>&#x002A;&#x002A;&#x002A;</sup><italic>p</italic>&#x2009;&#x003C;&#x2009;0.001.</p>
</caption>
<graphic xlink:href="fnins-18-1368641-g003.tif"/>
</fig>
<table-wrap position="float" id="tab2">
<label>Table 2</label>
<caption>
<p>ANOVA table for vocoded speech perception.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th/>
<th align="center" valign="top">Sum of square</th>
<th align="center" valign="top">df</th>
<th align="center" valign="top">Mean square</th>
<th align="center" valign="top">F</th>
<th align="center" valign="top"><italic>p</italic></th>
<th align="center" valign="top">&#x03B7;p<sup>2</sup></th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">Channel</td>
<td align="center" valign="top">2581.938</td>
<td align="center" valign="top">3</td>
<td align="center" valign="top">860.646</td>
<td align="center" valign="top">284.697</td>
<td align="center" valign="top">&#x003C;0.001</td>
<td align="center" valign="top">0.937</td>
</tr>
<tr>
<td align="left" valign="top">Residual</td>
<td align="center" valign="top">172.313</td>
<td align="center" valign="top">57</td>
<td align="center" valign="top">3.023</td>
<td/>
<td/>
<td/>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="sec13">
<label>3.2</label>
<title>Effects of spectral degradation on cortical representation</title>
<p><xref ref-type="fig" rid="fig4">Figure 4A</xref> displays the grand-averaged waveforms of the N2 component with waveforms elicited by standard stimuli (animal) in red and target (non-animal) in black. <xref ref-type="fig" rid="fig4">Figure 4B</xref> illustrates the grand-averaged waveforms and the difference between animal and non-animal stimuli for the P3b component. Shades represent the areas of N2 (light blue) and P3b (pink). RM ANOVA (four channel-band conditions) for N2 showed significant main effects of the number of channel bands on peak amplitude [<italic>F</italic><sub>(2.006, 38.117)</sub>&#x2009;=&#x2009;9.077, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001, <italic>&#x03B7;p<sup>2</sup></italic>&#x2009;=&#x2009;0.323] and peak latency [<italic>F</italic><sub>(3, 57)</sub>&#x2009;=&#x2009;26.642, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001, <italic>&#x03B7;p<sup>2</sup></italic>&#x2009;=&#x2009;0.584; <xref ref-type="fig" rid="fig4">Figure 4A</xref>; <xref ref-type="table" rid="tab3">Table 3</xref>]. In the comparisons of N2 peak amplitude based on the number of channel bands shown in <xref ref-type="fig" rid="fig5">Figure 5A</xref>, significant differences were observed between the conditions with 4 and 8 channel bands and the condition with 32 channels (all <italic>P<sub>Bonf</sub></italic>&#x2009;&#x003C;&#x2009;0.001). The peak amplitudes were 0.19&#x2009;&#x00B1;&#x2009;0.10, 0.21&#x2009;&#x00B1;&#x2009;0.08, 0.30&#x2009;&#x00B1;&#x2009;0.20, and 0.37&#x2009;&#x00B1;&#x2009;0.16&#x2009;&#x03BC;V for the conditions with 4, 8, 16, and 32 channel bands, respectively. Regarding the comparisons of N2 peak latency, as illustrated in <xref ref-type="fig" rid="fig5">Figure 5B</xref>, all pairs were significantly different from each other (all <italic>P<sub>Bonf</sub></italic>&#x2009;&#x003C;&#x2009;0.05) except for the conditions with 8 vs. 16 channel bands. The peak latencies were 504.38&#x2009;&#x00B1;&#x2009;63.43, 435.06&#x2009;&#x00B1;&#x2009;77.44, 398.75&#x2009;&#x00B1;&#x2009;63.12, ad 331.50&#x2009;&#x00B1;&#x2009;73.35&#x2009;ms for the conditions with 4, 8, 16, and 32 channel bands, respectively.</p>
<fig position="float" id="fig4">
<label>Figure 4</label>
<caption>
<p>Grand average waveforms of N2 <bold>(A)</bold> and P3b <bold>(B)</bold> components at each number of channel conditions. N2 were measured by averaging the signals in the frontocentral electrodes (Fz, FC1, FC2, and Cz) and P3b were measured in the parietal electrodes (CP1, CP2, P3, P4, and Pz). The area time windows determined from the current data were 280&#x2013;870&#x2009;msec for N2 and 280&#x2013;840&#x2009;msec for P3b components.</p>
</caption>
<graphic xlink:href="fnins-18-1368641-g004.tif"/>
</fig>
<table-wrap position="float" id="tab3">
<label>Table 3</label>
<caption>
<p>ANOVA table for the amplitude and latency in N2.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th/>
<th/>
<th align="center" valign="top">Sum of square</th>
<th align="center" valign="top">df</th>
<th align="center" valign="top">Mean square</th>
<th align="center" valign="top">F</th>
<th align="center" valign="top"><italic>p</italic></th>
<th align="center" valign="top">&#x03B7;p<sup>2</sup></th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top" rowspan="2">Amplitude</td>
<td align="left" valign="top">Channel</td>
<td align="center" valign="top">0.397</td>
<td align="center" valign="top">2.006</td>
<td align="center" valign="top">0.198</td>
<td align="center" valign="top">9.077</td>
<td align="center" valign="top">&#x003C;0.001</td>
<td align="center" valign="top">0.323</td>
</tr>
<tr>
<td align="left" valign="top">Residual</td>
<td align="center" valign="top">0.831</td>
<td align="center" valign="top">38.117</td>
<td align="center" valign="top">0.022</td>
<td/>
<td/>
<td/>
</tr>
<tr>
<td align="left" valign="top" rowspan="2">Latency</td>
<td align="left" valign="top">Channel</td>
<td align="center" valign="top">312064.902</td>
<td align="center" valign="top">3</td>
<td align="center" valign="top">104021.634</td>
<td align="center" valign="top">26.642</td>
<td align="center" valign="top">&#x003C;0.001</td>
<td align="center" valign="top">0.584</td>
</tr>
<tr>
<td align="left" valign="top">Residual</td>
<td align="center" valign="top">222549.551</td>
<td align="center" valign="top">57</td>
<td align="center" valign="top">3904.378</td>
<td/>
<td/>
<td/>
</tr>
</tbody>
</table>
</table-wrap>
<fig position="float" id="fig5">
<label>Figure 5</label>
<caption>
<p>Amplitude and latency comparisons of N2 and P3b between channel bands. Regarding the N2 peak amplitude <bold>(A)</bold>, there was a significant difference between 4 and 8 channel bands compared to 32 channels (all <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001), and for N2 peak latency <bold>(B)</bold>, all pairs showed significant differences (all <italic>p</italic>&#x2009;&#x003C;&#x2009;0.05). In terms of P3b amplitude comparisons, 4 channel bands significantly differed from all channels (all <italic>p</italic>&#x2009;&#x003C;&#x2009;0.01), while 8 channel bands differed from 32 channel bands (<italic>p</italic>&#x2009;=&#x2009;0.010) <bold>(C)</bold>. All pairs showed no significant differences for P3b latency <bold>(D)</bold>. <sup>&#x002A;</sup><italic>p</italic>&#x2009;&#x003C;&#x2009;0.05, <sup>&#x002A;&#x002A;</sup><italic>p</italic>&#x2009;&#x003C;&#x2009;0.01, <sup>&#x002A;&#x002A;&#x002A;</sup><italic>p</italic>&#x2009;&#x003C;&#x2009;0.001.</p>
</caption>
<graphic xlink:href="fnins-18-1368641-g005.tif"/>
</fig>
<p>RM ANOVA (four channel-band conditions) for P3b difference showed significant main effects of the number of channel bands on peak amplitude [<italic>F</italic><sub>(2.231, 42.391)</sub>&#x2009;=&#x2009;13.045, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001, <italic>&#x03B7;p<sup>2</sup></italic>&#x2009;=&#x2009;0.407] and peak latency [<italic>F</italic><sub>(3, 57)</sub>&#x2009;=&#x2009;2.968, <italic>p</italic>&#x2009;=&#x2009;0.039, <italic>&#x03B7;p<sup>2</sup></italic>&#x2009;=&#x2009;0.135; <xref ref-type="fig" rid="fig4">Figure 4B</xref>; <xref ref-type="table" rid="tab4">Table 4</xref>]. Regarding the comparisons of P3b amplitude in terms of the number of channel bands shown in <xref ref-type="fig" rid="fig5">Figure 5C</xref>, the condition with 4 channel bands differed significantly from the condition with all channels (all <italic>P<sub>Bonf</sub></italic>&#x2009;&#x003C;&#x2009;0.01), and the condition with 8 channel bands differed from the condition with 32 channel bands (<italic>P<sub>Bonf</sub></italic>&#x2009;=&#x2009;0.010). The peak amplitudes of P3b were 0.19&#x2009;&#x00B1;&#x2009;0.08, 0.26&#x2009;&#x00B1;&#x2009;0.08, 0.31&#x2009;&#x00B1;&#x2009;0.11, and 0.37&#x2009;&#x00B1;&#x2009;0.14 &#x03BC;V for the conditions with 4, 8, 16, and 32 channel bands, respectively. There were no significant differences in pairs in post-hoc comparisons for P3b latency. The peak latencies of P3b were 422.80&#x2009;&#x00B1;&#x2009;76.35, 403.05&#x2009;&#x00B1;&#x2009;59.41, 373.20&#x2009;&#x00B1;&#x2009;64.07, and 363.75&#x2009;&#x00B1;&#x2009;77.57&#x2009;ms for the conditions with 4, 8, 16, and 32 channel bands, respectively (<xref ref-type="fig" rid="fig5">Figure 5D</xref>).</p>
<table-wrap position="float" id="tab4">
<label>Table 4</label>
<caption>
<p>ANOVA table for the amplitude and latency in P3b.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th/>
<th/>
<th align="center" valign="top">Sum of square</th>
<th align="center" valign="top">df</th>
<th align="center" valign="top">Mean square</th>
<th align="center" valign="top">F</th>
<th align="center" valign="top"><italic>p</italic></th>
<th align="center" valign="top">&#x03B7;p<sup>2</sup></th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top" rowspan="2">Amplitude</td>
<td align="left" valign="top">Channel</td>
<td align="center" valign="top">0.343</td>
<td align="center" valign="top">2.231</td>
<td align="center" valign="top">0.154</td>
<td align="center" valign="top">13.045</td>
<td align="center" valign="top">&#x003C;0.001</td>
<td align="center" valign="top">0.407</td>
</tr>
<tr>
<td align="left" valign="top">Residual</td>
<td align="center" valign="top">0.499</td>
<td align="center" valign="top">42.391</td>
<td align="center" valign="top">0.012</td>
<td/>
<td/>
<td/>
</tr>
<tr>
<td align="left" valign="top" rowspan="2">Latency</td>
<td align="left" valign="top">Channel</td>
<td align="center" valign="top">44309.700</td>
<td align="center" valign="top">3</td>
<td align="center" valign="top">14769.900</td>
<td align="center" valign="top">2.968</td>
<td align="center" valign="top">0.039</td>
<td align="center" valign="top">0.135</td>
</tr>
<tr>
<td align="left" valign="top">Residual</td>
<td align="center" valign="top">283689.800</td>
<td align="center" valign="top">57</td>
<td align="center" valign="top">4977.014</td>
<td/>
<td/>
<td/>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="sec14">
<label>3.3</label>
<title>Correlation of neural response with behavioral data</title>
<p>We also determined correlations between vocoded speech perception and the neural response of N2 and P3b in terms of latency and amplitude. Significant correlations were found between the N2 peak amplitude/latency and the behavioral accuracy in vocoded speech perception (amplitude: <italic>r</italic>&#x2009;=&#x2009;0.428, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001; latency: <italic>r</italic>&#x2009;=&#x2009;&#x2212;0.599, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001; <xref ref-type="fig" rid="fig6">Figure 6</xref>). Similarly, there were significant correlations between the P3b peak amplitude/latency and the behavioral accuracy in vocoded speech perception (amplitude: <italic>r</italic>&#x2009;=&#x2009;0.531, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.001; latency: <italic>r</italic>&#x2009;=&#x2009;&#x2212;0.313, <italic>p</italic>&#x2009;=&#x2009;0.005; <xref ref-type="fig" rid="fig6">Figure 6</xref>).</p>
<fig position="float" id="fig6">
<label>Figure 6</label>
<caption>
<p>Correlation of neural response with behavioral data. There were significant correlations between the behavioral responses and N2 and P3b peak amplitudes/latencies (all <italic>p</italic>&#x2009;&#x003C;&#x2009;0.01).</p>
</caption>
<graphic xlink:href="fnins-18-1368641-g006.tif"/>
</fig>
</sec>
</sec>
<sec sec-type="discussion" id="sec15">
<label>4</label>
<title>Discussion</title>
<p>We found slowed N2 and P3b latencies, indicating delayed access to lexical information and semantic categorization in a small number of channel band conditions. Although neural responses were not significantly different in all pairs of channel bands, there was a consistent pattern of increasing amplitude and decreasing latency as the number of channel bands increased, aligning with the behavioral results. Behavioral tests indicated, as expected, that increasing the number of channel bands led to a stepwise improvement in the test of intelligibility, with all pairs showing significant differences. In addition, we found a strong correlation between N2 and P3b peak amplitudes/latencies and behavioral accuracy. This supports the effective representation of the channel band effect in a noise vocoder on N2 and P3b, confirming our hypothesis that spectral degradation with a noise vocoder significantly influences the neural representation of semantic processing.</p>
<p>Some studies showed the vocoder channel effect on P1-N1-P2 responses using CVC tokens (<xref ref-type="bibr" rid="ref16">Friesen et al., 2009</xref>) or the DISH-DITCH continuum (<xref ref-type="bibr" rid="ref2">Anderson et al., 2020</xref>). However, these studies compared the vocoded condition with the unprocessed condition and did not involve channel-specific multiple comparisons as in our current study. Moreover, a recent study found inconsistent changes in the P1, N1, and P2 components of ERP when listening to vocoded speech on 4 and 22 channels (<xref ref-type="bibr" rid="ref10">Dong and Gai, 2021</xref>).</p>
<p>Other studies have examined later ERPs (&#x003E;300&#x2009;ms post-stimuli) to assess cognitive spare capacity for measuring channel band effect (<xref ref-type="bibr" rid="ref36">Strauss et al., 2013</xref>; <xref ref-type="bibr" rid="ref4">Banellis et al., 2020</xref>; <xref ref-type="bibr" rid="ref19">Hunter, 2020</xref>). <xref ref-type="bibr" rid="ref4">Banellis et al. (2020)</xref> investigated the role of the P3a component and the impact of top-down expectations in a word-pair priming task involving degraded (noise-vocoded) speech. The findings suggest that expectations play a crucial role in the comprehension of degraded speech, influencing neural responses. <xref ref-type="bibr" rid="ref19">Hunter (2020)</xref> explored cognitive demand during listening to noise-vocoded spoken sentences, analyzing the impact of cognitive (memory) load and sentence predictability on electrophysiological measures, specifically the P300/late positive complex and N400. <xref ref-type="bibr" rid="ref36">Strauss et al. (2013)</xref> measured the N400 evoked by a congruent&#x2014;incongruent final object in a sentence paradigm using vocoders. In the 8-band condition, the N400 amplitude was attenuated, and the peak was significantly delayed by approximately 78&#x2009;ms compared with clean speech. They also demonstrated that the effect of spectral degradation on semantic processing involving expectation and context was more pronounced in ordinary speech than in vocoded speech. However, such congruent/incongruent semantic paradigms in sentences come with the limitation of being dependent on contextual cues, making it challenging to control for individuals&#x2019; education and cognitive abilities. It has been acknowledged that humans rely more on top-down processing when the spectral information in the speech signal is degraded (<xref ref-type="bibr" rid="ref34">Shannon et al., 1995</xref>; <xref ref-type="bibr" rid="ref7">Davis et al., 2005</xref>; <xref ref-type="bibr" rid="ref27">Obleser and Eisner, 2009</xref>; <xref ref-type="bibr" rid="ref31">Peelle and Davis, 2012</xref>).</p>
<p>To explore the impact of vocoder channel count on semantic processing without contextual cues, we utilized a one-syllable oddball paradigm, and measured the N2 and P3b components associated with purely acoustic-based semantic processing. Based on our results, we suggest that the N2 and P3b responses induced by a one-syllable task generated using a vocoder may serve as suitable objective measures for representing spectral degradation as a function of the number of channel bands. The earlier cortical potentials, between 100 and 200&#x2009;ms after the stimulus, are more related to the bottom-up perception of speech, and there is a possibility of a robust response even if the meaning of spectrally degraded speech cannot be understood. Therefore, the impact of spectral degradation in a noise vocoder could be more closely related to the semantic system than to the perception itself, especially when there is minimal redundancy in the available cues. Collectively, our results suggest that N2 and P3b responses, measuring the top-down mechanism of speech comprehension, would be useful tools for representing the effects of the number of channel bands.</p>
<p>Several studies have examined the neural representation of vocoded speech at the brainstem level. Using the frequency following responses (FFR), for example, neural facilitation was intimated as a function of the number of channel bands (<xref ref-type="bibr" rid="ref1">Ananthakrishnan et al., 2017</xref>) reported that, using FFR, the improvement in brainstem F0 magnitudes, phase-locked to the temporal envelope of the stimuli as the number of channels increased from 1 to 4 consistent with the behavioral performance. However, the F0 representation was followed by a plateau with 8 and 16 channels and then a degradation with 32 channels. Using FFR and cortical auditory-evoked potentials (<xref ref-type="bibr" rid="ref23">Kong et al., 2015</xref>), confirmed that attention modulates EEG entrainment to the speech envelope and that the neural entrainment measured using correlation coefficients between the N1 response and speech envelope increased as a function of the number of channel bands.</p>
<p>A limitation of this study was the relatively small size of the study sample. However, we endeavored to control the age of the participants along with their duration of education and attention level throughout the tests. Because we controlled the age of the listeners (young participants), we could not generalize our conclusion to a wider population, such as the elderly, in the current study. Another limitation is that in the analysis of the EEG response, both correct and incorrect responses were used without excluding the incorrect ones due to the very low correct response rate. Although excluding trials with an incorrect behavioral response could ensure a more accurate interpretation of the EEG response, we were concerned that deleting too much data would decrease the reliability of the overall dataset. Future research could consider expanding participants to the various age groups. Furthermore, it would be necessary to conduct studies measuring cortical responses in listeners with hearing impairment characterized by sparse spectral resolution.</p>
<p>Our data enhance understanding of how spectral information influences cortical speech processing, and they have implications for developing advanced algorithms in hearing aids for individuals with hearing impairments or degraded auditory input. Moreover, a better understanding of the impact of sparse spectral information on speech intelligibility and neural representation could provide valuable insights for designing public places such as auditoriums or transportation hubs.</p>
</sec>
<sec sec-type="conclusions" id="sec16">
<label>5</label>
<title>Conclusion</title>
<p>Our study has demonstrated that the degree of spectral richness plays a crucial role in both speech perception and neural responses. We particularly adopted a one-syllable paradigm to elicit N2 and P3b responses, a higher-order cognitive process, while minimizing contextual cues and controlling the education and the attention level across participants. The N2 and P3b responses proved to be highly sensitive to the effects of spectral degradation, surpassing its sensitivity to speech perception.</p>
</sec>
<sec sec-type="data-availability" id="sec17">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/supplementary material, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec sec-type="ethics-statement" id="sec18">
<title>Ethics statement</title>
<p>The studies involving humans were approved by Review Board of Nowon Eulji Medical Center. The studies were conducted in accordance with the local legislation and institutional requirements. The participants provided their written informed consent to participate in this study.</p>
</sec>
<sec sec-type="author-contributions" id="sec19">
<title>Author contributions</title>
<p>HC: Methodology, Writing &#x2013; original draft, Data curation, Investigation. J-SK: Data curation, Formal analysis, Methodology, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. JW: Methodology, Writing &#x2013; review &#x0026; editing. HS: Conceptualization, Funding acquisition, Methodology, Supervision, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing.</p>
</sec>
</body>
<back>
<sec sec-type="funding-information" id="sec20">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research, authorship, and/or publication of this article. This research was supported by Basic Science Research Program through the National Research Foundation of Korea (NRF) funded by the Ministry of Education (2020R1I1A3071587).</p>
</sec>
<ack>
<p>Hyunsook Jang at Hallym University provided the monosyllabic word lists.</p>
</ack>
<sec sec-type="COI-statement" id="sec21">
<title>Conflict of interest</title>
<p>Author JW was employed by company Hyman, Phelps and McNamara, P.C.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="sec100" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="ref1">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ananthakrishnan</surname> <given-names>S.</given-names></name> <name><surname>Luo</surname> <given-names>X.</given-names></name> <name><surname>Krishnan</surname> <given-names>A.</given-names></name></person-group> (<year>2017</year>). <article-title>Human frequency following responses to vocoded speech</article-title>. <source>Ear Hear.</source> <volume>38</volume>:<fpage>e256</fpage>, &#x2013;<lpage>e267</lpage>. doi: <pub-id pub-id-type="doi">10.1097/AUD.0000000000000432</pub-id>, PMID: <pub-id pub-id-type="pmid">28362674</pub-id></citation>
</ref>
<ref id="ref2">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Anderson</surname> <given-names>S.</given-names></name> <name><surname>Roque</surname> <given-names>L.</given-names></name> <name><surname>Gaskins</surname> <given-names>C. R.</given-names></name> <name><surname>Gordon-Salant</surname> <given-names>S.</given-names></name> <name><surname>Goupell</surname> <given-names>M. J.</given-names></name></person-group> (<year>2020</year>). <article-title>Age-related compensation mechanism revealed in the cortical representation of degraded speech</article-title>. <source>J. Assoc. Res. Otolaryngol.</source> <volume>21</volume>, <fpage>373</fpage>&#x2013;<lpage>391</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10162-020-00753-4</pub-id>, PMID: <pub-id pub-id-type="pmid">32643075</pub-id></citation>
</ref>
<ref id="ref3">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bae</surname> <given-names>E. B.</given-names></name> <name><surname>Jang</surname> <given-names>H.</given-names></name> <name><surname>Shim</surname> <given-names>H. J.</given-names></name></person-group> (<year>2022</year>). <article-title>Enhanced dichotic listening and temporal sequencing ability in early-blind individuals</article-title>. <source>Front. Psychol.</source> <volume>13</volume>:<fpage>840541</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2022.840541</pub-id>, PMID: <pub-id pub-id-type="pmid">35619788</pub-id></citation>
</ref>
<ref id="ref4">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Banellis</surname> <given-names>L.</given-names></name> <name><surname>Sokoliuk</surname> <given-names>R.</given-names></name> <name><surname>Wild</surname> <given-names>C. J.</given-names></name> <name><surname>Bowman</surname> <given-names>H.</given-names></name> <name><surname>Cruse</surname> <given-names>D.</given-names></name></person-group> (<year>2020</year>). <article-title>Event-related potentials reflect prediction errors and pop-out during comprehension of degraded speech</article-title>. <source>Neurosci Conscious</source> <volume>2020</volume>:<fpage>niaa022</fpage>. doi: <pub-id pub-id-type="doi">10.1093/nc/niaa022</pub-id></citation>
</ref>
<ref id="ref5">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Beynon</surname> <given-names>A.</given-names></name> <name><surname>Snik</surname> <given-names>A.</given-names></name> <name><surname>Stegeman</surname> <given-names>D.</given-names></name> <name><surname>Van den Broek</surname> <given-names>P.</given-names></name></person-group> (<year>2005</year>). <article-title>Discrimination of speech sound contrasts determined with behavioral tests and event-related potentials in cochlear implant recipients</article-title>. <source>J. Am. Acad. Audiol.</source> <volume>16</volume>, <fpage>042</fpage>&#x2013;<lpage>053</lpage>. doi: <pub-id pub-id-type="doi">10.3766/jaaa.16.1.5</pub-id></citation>
</ref>
<ref id="ref6">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Davis</surname> <given-names>M. H.</given-names></name> <name><surname>Johnsrude</surname> <given-names>I. S.</given-names></name></person-group> (<year>2003</year>). <article-title>Hierarchical processing in spoken language comprehension</article-title>. <source>J. Neurosci.</source> <volume>23</volume>, <fpage>3423</fpage>&#x2013;<lpage>3431</lpage>. doi: <pub-id pub-id-type="doi">10.1523/JNEUROSCI.23-08-03423.2003</pub-id>, PMID: <pub-id pub-id-type="pmid">12716950</pub-id></citation>
</ref>
<ref id="ref7">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Davis</surname> <given-names>M. H.</given-names></name> <name><surname>Johnsrude</surname> <given-names>I. S.</given-names></name> <name><surname>Hervais-Adelman</surname> <given-names>A.</given-names></name> <name><surname>Taylor</surname> <given-names>K.</given-names></name> <name><surname>McGettigan</surname> <given-names>C.</given-names></name></person-group> (<year>2005</year>). <article-title>Lexical information drives perceptual learning of distorted speech: evidence from the comprehension of noise-vocoded sentences</article-title>. <source>J. Exp. Psychol. Gen.</source> <volume>134</volume>, <fpage>222</fpage>&#x2013;<lpage>241</lpage>. doi: <pub-id pub-id-type="doi">10.1037/0096-3445.134.2.222</pub-id>, PMID: <pub-id pub-id-type="pmid">15869347</pub-id></citation>
</ref>
<ref id="ref8">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Deacon</surname> <given-names>D.</given-names></name> <name><surname>Breton</surname> <given-names>F.</given-names></name> <name><surname>Ritter</surname> <given-names>W.</given-names></name> <name><surname>Vaughan</surname> <given-names>H. G.</given-names> <suffix>Jr.</suffix></name></person-group> (<year>1991</year>). <article-title>The relationship between N2 and N400: scalp distribution, stimulus probability, and task relevance</article-title>. <source>Psychophysiology</source> <volume>28</volume>, <fpage>185</fpage>&#x2013;<lpage>200</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1469-8986.1991.tb00411.x</pub-id>, PMID: <pub-id pub-id-type="pmid">1946885</pub-id></citation>
</ref>
<ref id="ref9">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Delorme</surname> <given-names>A.</given-names></name> <name><surname>Makeig</surname> <given-names>S.</given-names></name></person-group> (<year>2004</year>). <article-title>EEGLAB: an open source toolbox for analysis of single-trial EEG dynamics including independent component analysis</article-title>. <source>J. Neurosci. Methods</source> <volume>134</volume>, <fpage>9</fpage>&#x2013;<lpage>21</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.jneumeth.2003.10.009</pub-id>, PMID: <pub-id pub-id-type="pmid">15102499</pub-id></citation>
</ref>
<ref id="ref10">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dong</surname> <given-names>Y.</given-names></name> <name><surname>Gai</surname> <given-names>Y.</given-names></name></person-group> (<year>2021</year>). <article-title>Speech perception with noise Vocoding and background noise: an EEG and behavioral study</article-title>. <source>J. Assoc. Res. Otolaryngol.</source> <volume>22</volume>, <fpage>349</fpage>&#x2013;<lpage>363</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10162-021-00787-2</pub-id>, PMID: <pub-id pub-id-type="pmid">33851289</pub-id></citation>
</ref>
<ref id="ref11">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dorman</surname> <given-names>M. F.</given-names></name> <name><surname>Loizou</surname> <given-names>P. C.</given-names></name> <name><surname>Rainey</surname> <given-names>D.</given-names></name></person-group> (<year>1997</year>). <article-title>Speech intelligibility as a function of the number of channels of stimulation for signal processors using sine-wave and noise-band outputs</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>102</volume>, <fpage>2403</fpage>&#x2013;<lpage>2411</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.419603</pub-id>, PMID: <pub-id pub-id-type="pmid">9348698</pub-id></citation>
</ref>
<ref id="ref12">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Evans</surname> <given-names>S.</given-names></name> <name><surname>Kyong</surname> <given-names>J.</given-names></name> <name><surname>Rosen</surname> <given-names>S.</given-names></name> <name><surname>Golestani</surname> <given-names>N.</given-names></name> <name><surname>Warren</surname> <given-names>J.</given-names></name> <name><surname>McGettigan</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2014</year>). <article-title>The pathways for intelligible speech: multivariate and univariate perspectives</article-title>. <source>Cereb. Cortex</source> <volume>24</volume>, <fpage>2350</fpage>&#x2013;<lpage>2361</lpage>. doi: <pub-id pub-id-type="doi">10.1093/cercor/bht083</pub-id>, PMID: <pub-id pub-id-type="pmid">23585519</pub-id></citation>
</ref>
<ref id="ref1001">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Faul</surname> <given-names>F.</given-names></name> <name><surname>Erdfelder</surname> <given-names>E.</given-names></name> <name><surname>Lang</surname> <given-names>A.-G.</given-names></name> <name><surname>Buchner</surname> <given-names>T.</given-names></name></person-group> (<year>2007</year>). <article-title>G&#x002A; Power 3: A flexible statistical power analysis program for the social, behavioral, and biomedical sciences</article-title>. <source>Behav. Res. Methods.</source> <volume>39</volume>, <fpage>175</fpage>&#x2013;<lpage>191</lpage>.</citation>
</ref>
<ref id="ref13">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Faulkner</surname> <given-names>A.</given-names></name> <name><surname>Rosen</surname> <given-names>S.</given-names></name> <name><surname>Green</surname> <given-names>T.</given-names></name></person-group> (<year>2012</year>). <article-title>Comparing live to recorded speech in training the perception of spectrally shifted noise-vocoded speech</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>132</volume>:<fpage>EL336-EL342</fpage>. doi: <pub-id pub-id-type="doi">10.1121/1.4754432</pub-id>, PMID: <pub-id pub-id-type="pmid">23039574</pub-id></citation>
</ref>
<ref id="ref14">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Finke</surname> <given-names>M.</given-names></name> <name><surname>B&#x00FC;chner</surname> <given-names>A.</given-names></name> <name><surname>Ruigendijk</surname> <given-names>E.</given-names></name> <name><surname>Meyer</surname> <given-names>M.</given-names></name> <name><surname>Sandmann</surname> <given-names>P.</given-names></name></person-group> (<year>2016</year>). <article-title>On the relationship between auditory cognition and speech intelligibility in cochlear implant users: an ERP study</article-title>. <source>Neuropsychologia</source> <volume>87</volume>, <fpage>169</fpage>&#x2013;<lpage>181</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.neuropsychologia.2016.05.019</pub-id>, PMID: <pub-id pub-id-type="pmid">27212057</pub-id></citation>
</ref>
<ref id="ref15">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Folstein</surname> <given-names>J. R.</given-names></name> <name><surname>Van Petten</surname> <given-names>C.</given-names></name></person-group> (<year>2008</year>). <article-title>Influence of cognitive control and mismatch on the N2 component of the ERP: a review</article-title>. <source>Psychophysiology</source> <volume>45</volume>, <fpage>152</fpage>&#x2013;<lpage>170</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1469-8986.2007.00602.x</pub-id>, PMID: <pub-id pub-id-type="pmid">17850238</pub-id></citation>
</ref>
<ref id="ref16">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Friesen</surname> <given-names>L.</given-names></name> <name><surname>Tremblay</surname> <given-names>K.</given-names></name> <name><surname>Rohila</surname> <given-names>N.</given-names></name> <name><surname>Wright</surname> <given-names>R.</given-names></name> <name><surname>Shannon</surname> <given-names>R.</given-names></name> <name><surname>Ba&#x015F;kent</surname> <given-names>D.</given-names></name> <etal/></person-group>. (<year>2009</year>). <article-title>Evoked cortical activity and speech recognition as a function of the number of simulated cochlear implant channels</article-title>. <source>Clin. Neurophysiol.</source> <volume>120</volume>, <fpage>776</fpage>&#x2013;<lpage>782</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.clinph.2009.01.008</pub-id>, PMID: <pub-id pub-id-type="pmid">19250865</pub-id></citation>
</ref>
<ref id="ref17">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Henkin</surname> <given-names>Y.</given-names></name> <name><surname>Yaar-Soffer</surname> <given-names>Y.</given-names></name> <name><surname>Givon</surname> <given-names>L.</given-names></name> <name><surname>Hildesheimer</surname> <given-names>M.</given-names></name></person-group> (<year>2015</year>). <article-title>Hearing with two ears: evidence for cortical binaural interaction during auditory processing</article-title>. <source>J. Am. Acad. Audiol.</source> <volume>26</volume>, <fpage>384</fpage>&#x2013;<lpage>392</lpage>. doi: <pub-id pub-id-type="doi">10.3766/jaaa.26.4.6</pub-id></citation>
</ref>
<ref id="ref18">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hervais-Adelman</surname> <given-names>A.</given-names></name> <name><surname>Davis</surname> <given-names>M. H.</given-names></name> <name><surname>Johnsrude</surname> <given-names>I. S.</given-names></name> <name><surname>Carlyon</surname> <given-names>R. P.</given-names></name></person-group> (<year>2008</year>). <article-title>Perceptual learning of noise vocoded words: effects of feedback and lexicality</article-title>. <source>J. Exp. Psychol. Hum. Percept. Perform.</source> <volume>34</volume>, <fpage>460</fpage>&#x2013;<lpage>474</lpage>. doi: <pub-id pub-id-type="doi">10.1037/0096-1523.34.2.460</pub-id>, PMID: <pub-id pub-id-type="pmid">18377182</pub-id></citation>
</ref>
<ref id="ref19">
<citation citation-type="journal"><person-group person-group-type="author">
<name><surname>Hunter</surname> <given-names>C. R.</given-names></name>
</person-group> (<year>2020</year>). <article-title>Tracking cognitive spare capacity during speech perception with EEG/ERP: effects of cognitive load and sentence predictability</article-title>. <source>Ear Hear.</source> <volume>41</volume>, <fpage>1144</fpage>&#x2013;<lpage>1157</lpage>. doi: <pub-id pub-id-type="doi">10.1097/AUD.0000000000000856</pub-id>, PMID: <pub-id pub-id-type="pmid">32282402</pub-id></citation>
</ref>
<ref id="ref20">
<citation citation-type="journal"><person-group person-group-type="author">
<name><surname>Jasper</surname> <given-names>H.</given-names></name>
</person-group> (<year>1958</year>). <article-title>Report of the committee on methods of clinical examination in electroencephalography: 1957</article-title>. <source>Electroencephalogr. Clin. Neurophysiol.</source> <volume>10</volume>, <fpage>370</fpage>&#x2013;<lpage>375</lpage>.</citation>
</ref>
<ref id="ref21">
<citation citation-type="journal"><person-group person-group-type="author">
<name><surname>Johnson</surname> <given-names>R.</given-names></name>
</person-group> (<year>1988</year>). <article-title>The amplitude of the P300 component of the event-related potential: review and synthesis</article-title>. <source>Adv Psychophysiol</source> <volume>3</volume>, <fpage>69</fpage>&#x2013;<lpage>137</lpage>.</citation>
</ref>
<ref id="ref22">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jung</surname> <given-names>T.-P.</given-names></name> <name><surname>Makeig</surname> <given-names>S.</given-names></name> <name><surname>Humphries</surname> <given-names>C.</given-names></name> <name><surname>Lee</surname> <given-names>T.-W.</given-names></name> <name><surname>Mckeown</surname> <given-names>M. J.</given-names></name> <name><surname>Iragui</surname> <given-names>V.</given-names></name> <etal/></person-group>. (<year>2000</year>). <article-title>Removing electroencephalographic artifacts by blind source separation</article-title>. <source>Psychophysiology</source> <volume>37</volume>, <fpage>163</fpage>&#x2013;<lpage>178</lpage>. doi: <pub-id pub-id-type="doi">10.1111/1469-8986.3720163</pub-id>, PMID: <pub-id pub-id-type="pmid">10731767</pub-id></citation>
</ref>
<ref id="ref23">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kong</surname> <given-names>Y.-Y.</given-names></name> <name><surname>Somarowthu</surname> <given-names>A.</given-names></name> <name><surname>Ding</surname> <given-names>N.</given-names></name></person-group> (<year>2015</year>). <article-title>Effects of spectral degradation on attentional modulation of cortical auditory responses to continuous speech</article-title>. <source>J. Assoc. Res. Otolaryngol.</source> <volume>16</volume>, <fpage>783</fpage>&#x2013;<lpage>796</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10162-015-0540-x</pub-id>, PMID: <pub-id pub-id-type="pmid">26362546</pub-id></citation>
</ref>
<ref id="ref24">
<citation citation-type="book"><person-group person-group-type="author">
<name><surname>Luck</surname> <given-names>S. J.</given-names></name>
</person-group> (<year>2014</year>). <source>An introduction to the event-related potential technique</source>. Cambrigde: MA: <publisher-name>MIT press</publisher-name>.</citation>
</ref>
<ref id="ref25">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Makeig</surname> <given-names>S.</given-names></name> <name><surname>Jung</surname> <given-names>T.-P.</given-names></name> <name><surname>Bell</surname> <given-names>A. J.</given-names></name> <name><surname>Ghahremani</surname> <given-names>D.</given-names></name> <name><surname>Sejnowski</surname> <given-names>T. J.</given-names></name></person-group> (<year>1997</year>). <article-title>Blind separation of auditory event-related brain responses into independent components</article-title>. <source>Proc. Natl. Acad. Sci.</source> <volume>94</volume>, <fpage>10979</fpage>&#x2013;<lpage>10984</lpage>. doi: <pub-id pub-id-type="doi">10.1073/pnas.94.20.10979</pub-id>, PMID: <pub-id pub-id-type="pmid">9380745</pub-id></citation>
</ref>
<ref id="ref26">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Narain</surname> <given-names>C.</given-names></name> <name><surname>Scott</surname> <given-names>S. K.</given-names></name> <name><surname>Wise</surname> <given-names>R. J.</given-names></name> <name><surname>Rosen</surname> <given-names>S.</given-names></name> <name><surname>Leff</surname> <given-names>A.</given-names></name> <name><surname>Iversen</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2003</year>). <article-title>Defining a left-lateralized response specific to intelligible speech using fMRI</article-title>. <source>Cereb. Cortex</source> <volume>13</volume>, <fpage>1362</fpage>&#x2013;<lpage>1368</lpage>. doi: <pub-id pub-id-type="doi">10.1093/cercor/bhg083</pub-id>, PMID: <pub-id pub-id-type="pmid">14615301</pub-id></citation>
</ref>
<ref id="ref27">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Obleser</surname> <given-names>J.</given-names></name> <name><surname>Eisner</surname> <given-names>F.</given-names></name></person-group> (<year>2009</year>). <article-title>Pre-lexical abstraction of speech in the auditory cortex</article-title>. <source>Trends Cogn. Sci.</source> <volume>13</volume>, <fpage>14</fpage>&#x2013;<lpage>19</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.tics.2008.09.005</pub-id>, PMID: <pub-id pub-id-type="pmid">19070534</pub-id></citation>
</ref>
<ref id="ref28">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Obleser</surname> <given-names>J.</given-names></name> <name><surname>Wise</surname> <given-names>R. J.</given-names></name> <name><surname>Dresner</surname> <given-names>M. A.</given-names></name> <name><surname>Scott</surname> <given-names>S. K.</given-names></name></person-group> (<year>2007</year>). <article-title>Functional integration across brain regions improves speech perception under adverse listening conditions</article-title>. <source>J. Neurosci.</source> <volume>27</volume>, <fpage>2283</fpage>&#x2013;<lpage>2289</lpage>. doi: <pub-id pub-id-type="doi">10.1523/JNEUROSCI.4663-06.2007</pub-id>, PMID: <pub-id pub-id-type="pmid">17329425</pub-id></citation>
</ref>
<ref id="ref29">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Oostenveld</surname> <given-names>R.</given-names></name> <name><surname>Fries</surname> <given-names>P.</given-names></name> <name><surname>Maris</surname> <given-names>E.</given-names></name> <name><surname>Schoffelen</surname> <given-names>J.-M.</given-names></name></person-group> (<year>2011</year>). <article-title>FieldTrip: open source software for advanced analysis of MEG, EEG, and invasive electrophysiological data</article-title>. <source>Comput. Intell. Neurosci.</source> <volume>2011</volume>, <fpage>1</fpage>&#x2013;<lpage>9</lpage>. doi: <pub-id pub-id-type="doi">10.1155/2011/156869</pub-id></citation>
</ref>
<ref id="ref30">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pals</surname> <given-names>C.</given-names></name> <name><surname>Sarampalis</surname> <given-names>A.</given-names></name> <name><surname>Beynon</surname> <given-names>A.</given-names></name> <name><surname>Stainsby</surname> <given-names>T.</given-names></name> <name><surname>Ba&#x015F;kent</surname> <given-names>D.</given-names></name></person-group> (<year>2020</year>). <article-title>Effect of spectral channels on speech recognition, comprehension, and listening effort in cochlear-implant users</article-title>. <source>Trends Hear</source> <volume>24</volume>:<fpage>233121652090461</fpage>. doi: <pub-id pub-id-type="doi">10.1177/2331216520904617</pub-id></citation>
</ref>
<ref id="ref31">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Peelle</surname> <given-names>J. E.</given-names></name> <name><surname>Davis</surname> <given-names>M. H.</given-names></name></person-group> (<year>2012</year>). <article-title>Neural oscillations carry speech rhythm through to comprehension</article-title>. <source>Front. Psychol.</source> <volume>3</volume>:<fpage>320</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2012.00320</pub-id></citation>
</ref>
<ref id="ref32">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Perrin</surname> <given-names>F.</given-names></name> <name><surname>Pernier</surname> <given-names>J.</given-names></name> <name><surname>Bertrand</surname> <given-names>O.</given-names></name> <name><surname>Echallier</surname> <given-names>J. F.</given-names></name></person-group> (<year>1989</year>). <article-title>Spherical splines for scalp potential and current density mapping</article-title>. <source>Electroencephalogr. Clin. Neurophysiol.</source> <volume>72</volume>, <fpage>184</fpage>&#x2013;<lpage>187</lpage>. doi: <pub-id pub-id-type="doi">10.1016/0013-4694(89)90180-6</pub-id>, PMID: <pub-id pub-id-type="pmid">2464490</pub-id></citation>
</ref>
<ref id="ref33">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schmitt</surname> <given-names>B. M.</given-names></name> <name><surname>M&#x00FC;nte</surname> <given-names>T. F.</given-names></name> <name><surname>Kutas</surname> <given-names>M.</given-names></name></person-group> (<year>2000</year>). <article-title>Electrophysiological estimates of the time course of semantic and phonological encoding during implicit picture naming</article-title>. <source>Psychophysiology</source> <volume>37</volume>, <fpage>473</fpage>&#x2013;<lpage>484</lpage>. doi: <pub-id pub-id-type="doi">10.1111/1469-8986.3740473</pub-id>, PMID: <pub-id pub-id-type="pmid">10934906</pub-id></citation>
</ref>
<ref id="ref34">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shannon</surname> <given-names>R. V.</given-names></name> <name><surname>Zeng</surname> <given-names>F.-G.</given-names></name> <name><surname>Kamath</surname> <given-names>V.</given-names></name> <name><surname>Wygonski</surname> <given-names>J.</given-names></name> <name><surname>Ekelid</surname> <given-names>M.</given-names></name></person-group> (<year>1995</year>). <article-title>Speech recognition with primarily temporal cues</article-title>. <source>Science</source> <volume>270</volume>, <fpage>303</fpage>&#x2013;<lpage>304</lpage>. doi: <pub-id pub-id-type="doi">10.1126/science.270.5234.303</pub-id></citation>
</ref>
<ref id="ref35">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Souza</surname> <given-names>P.</given-names></name> <name><surname>Rosen</surname> <given-names>S.</given-names></name></person-group> (<year>2009</year>). <article-title>Effects of envelope bandwidth on the intelligibility of sine-and noise-vocoded speech</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>126</volume>, <fpage>792</fpage>&#x2013;<lpage>805</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.3158835</pub-id>, PMID: <pub-id pub-id-type="pmid">19640044</pub-id></citation>
</ref>
<ref id="ref36">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Strauss</surname> <given-names>A.</given-names></name> <name><surname>Kotz</surname> <given-names>S. A.</given-names></name> <name><surname>Obleser</surname> <given-names>J.</given-names></name></person-group> (<year>2013</year>). <article-title>Narrowed expectancies under degraded speech: revisiting the N400</article-title>. <source>J. Cogn. Neurosci.</source> <volume>25</volume>, <fpage>1383</fpage>&#x2013;<lpage>1395</lpage>. doi: <pub-id pub-id-type="doi">10.1162/jocn_a_00389</pub-id>, PMID: <pub-id pub-id-type="pmid">23489145</pub-id></citation>
</ref>
<ref id="ref37">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Van den Brink</surname> <given-names>D.</given-names></name> <name><surname>Hagoort</surname> <given-names>P.</given-names></name></person-group> (<year>2004</year>). <article-title>The influence of semantic and syntactic context constraints on lexical selection and integration in spoken-word comprehension as revealed by ERPs</article-title>. <source>J. Cogn. Neurosci.</source> <volume>16</volume>, <fpage>1068</fpage>&#x2013;<lpage>1084</lpage>. doi: <pub-id pub-id-type="doi">10.1162/0898929041502670</pub-id></citation>
</ref>
<ref id="ref38">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Volpe</surname> <given-names>U.</given-names></name> <name><surname>Mucci</surname> <given-names>A.</given-names></name> <name><surname>Bucci</surname> <given-names>P.</given-names></name> <name><surname>Merlotti</surname> <given-names>E.</given-names></name> <name><surname>Galderisi</surname> <given-names>S.</given-names></name> <name><surname>Maj</surname> <given-names>M.</given-names></name></person-group> (<year>2007</year>). <article-title>The cortical generators of P3a and P3b: a LORETA study</article-title>. <source>Brain Res. Bull.</source> <volume>73</volume>, <fpage>220</fpage>&#x2013;<lpage>230</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.brainresbull.2007.03.003</pub-id>, PMID: <pub-id pub-id-type="pmid">17562387</pub-id></citation>
</ref>
<ref id="ref39">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Voola</surname> <given-names>M.</given-names></name> <name><surname>Nguyen</surname> <given-names>A. T.</given-names></name> <name><surname>Marinovic</surname> <given-names>W.</given-names></name> <name><surname>Rajan</surname> <given-names>G.</given-names></name> <name><surname>Tavora-Vieira</surname> <given-names>D.</given-names></name></person-group> (<year>2022</year>). <article-title>Odd-even oddball task: evaluating event-related potentials during word discrimination compared to speech-token and tone discrimination</article-title>. <source>Front. Neurosci.</source> <volume>16</volume>:<fpage>983498</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fnins.2022.983498</pub-id>, PMID: <pub-id pub-id-type="pmid">36312013</pub-id></citation>
</ref>
<ref id="ref40">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Voola</surname> <given-names>M.</given-names></name> <name><surname>Wedekind</surname> <given-names>A.</given-names></name> <name><surname>Nguyen</surname> <given-names>A. T.</given-names></name> <name><surname>Marinovic</surname> <given-names>W.</given-names></name> <name><surname>Rajan</surname> <given-names>G.</given-names></name> <name><surname>Tavora-Vieira</surname> <given-names>D.</given-names></name></person-group> (<year>2023</year>). <article-title>Event-related potentials of single-sided deaf Cochlear implant users: using a semantic oddball paradigm in noise</article-title>. <source>Audiol Neurotol</source> <volume>28</volume>, <fpage>280</fpage>&#x2013;<lpage>293</lpage>. doi: <pub-id pub-id-type="doi">10.1159/000529485</pub-id>, PMID: <pub-id pub-id-type="pmid">36940674</pub-id></citation>
</ref>
<ref id="ref41">
<citation citation-type="other"><person-group person-group-type="author"><name><surname>Xu</surname> <given-names>D.</given-names></name> <name><surname>Zheng</surname> <given-names>D.</given-names></name> <name><surname>Chen</surname> <given-names>F.</given-names></name></person-group> (<year>2019</year>). "Studying the effect of carrier type on the perception of vocoded stimuli via mismatch negativity", In <italic>2019 41st Annual International Conference of the IEEE Engineering in Medicine and Biology Society (EMBC)</italic>: IEEE, 3167&#x2013;3170.</citation>
</ref>
</ref-list>
</back>
</article>