<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Hum. Neurosci.</journal-id>
<journal-title>Frontiers in Human Neuroscience</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Hum. Neurosci.</abbrev-journal-title>
<issn pub-type="epub">1662-5161</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fnhum.2017.00034</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Neuroscience</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Selective Attention Enhances Beta-Band Cortical Oscillation to Speech under &#x201C;Cocktail-Party&#x201D; Listening Conditions</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name><surname>Gao</surname> <given-names>Yayue</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x2020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/378616/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Wang</surname> <given-names>Qian</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn002"><sup>&#x2020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/367706/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Ding</surname> <given-names>Yu</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/408241/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Wang</surname> <given-names>Changming</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/283165/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Li</surname> <given-names>Haifeng</given-names></name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/313146/overview"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Wu</surname> <given-names>Xihong</given-names></name>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
<xref ref-type="aff" rid="aff6"><sup>6</sup></xref>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Qu</surname> <given-names>Tianshu</given-names></name>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
<xref ref-type="aff" rid="aff6"><sup>6</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/384655/overview"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name><surname>Li</surname> <given-names>Liang</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/203218/overview"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Beijing Key Laboratory of Behavior and Mental Health, School of Psychological and Cognitive Sciences, Peking University</institution> <country>Beijing, China</country></aff>
<aff id="aff2"><sup>2</sup><institution>Beijing Anding Hospital, Capital Medical University</institution> <country>Beijing, China</country></aff>
<aff id="aff3"><sup>3</sup><institution>Beijing Institute for Brain Disorders, Capital Medical University</institution> <country>Beijing, China</country></aff>
<aff id="aff4"><sup>4</sup><institution>School of Computer Science and Technology, Harbin Institute of Technology</institution> <country>Harbin, China</country></aff>
<aff id="aff5"><sup>5</sup><institution>Department of Machine Intelligence, Peking University</institution> <country>Beijing, China</country></aff>
<aff id="aff6"><sup>6</sup><institution>Key Laboratory on Machine Perception &#x2013; Ministry of Education, Speech and Hearing Research Center, Peking University</institution> <country>Beijing, China</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: <italic>Xing Tian, New York University Shanghai, China</italic></p></fn>
<fn fn-type="edited-by"><p>Reviewed by: <italic>Dan Zhang, Tsinghua University, China; Nai Ding, Zhejiang University, China</italic></p></fn>
<fn fn-type="corresp" id="fn001"><p>&#x002A;Correspondence: <italic>Tianshu Qu, <email>qutianshu@pku.edu.cn</email> Liang Li, <email>liangli@pku.edu.cn</email></italic></p></fn>
<fn fn-type="other" id="fn002"><p><sup>&#x2020;</sup><italic>These authors are co-first authors.</italic></p></fn>
</author-notes>
<pub-date pub-type="epub">
<day>10</day>
<month>02</month>
<year>2017</year>
</pub-date>
<pub-date pub-type="collection">
<year>2017</year>
</pub-date>
<volume>11</volume>
<elocation-id>34</elocation-id>
<history>
<date date-type="received">
<day>09</day>
<month>09</month>
<year>2016</year>
</date>
<date date-type="accepted">
<day>16</day>
<month>01</month>
<year>2017</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2017 Gao, Wang, Ding, Wang, Li, Wu, Qu and Li.</copyright-statement>
<copyright-year>2017</copyright-year>
<copyright-holder>Gao, Wang, Ding, Wang, Li, Wu, Qu and Li</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) or licensor are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<p>Human listeners are able to selectively attend to target speech in a noisy environment with multiple-people talking. Using recordings of scalp electroencephalogram (EEG), this study investigated how selective attention facilitates the cortical representation of target speech under a simulated &#x201C;cocktail-party&#x201D; listening condition with speech-on-speech masking. The result shows that the cortical representation of target-speech signals under the multiple-people talking condition was specifically improved by selective attention relative to the non-selective-attention listening condition, and the beta-band activity was most strongly modulated by selective attention. Moreover, measured with the Granger Causality value, selective attention to the single target speech in the mixed-speech complex enhanced the following four causal connectivities for the beta-band oscillation: the ones (1) from site FT7 to the right motor area, (2) from the left frontal area to the right motor area, (3) from the central frontal area to the right motor area, and (4) from the central frontal area to the right frontal area. However, the selective-attention-induced change in beta-band causal connectivity from the central frontal area to the right motor area, but not other beta-band causal connectivities, was significantly correlated with the selective-attention-induced change in the cortical beta-band representation of target speech. These findings suggest that under the &#x201C;cocktail-party&#x201D; listening condition, the beta-band oscillation in EEGs to target speech is specifically facilitated by selective attention to the target speech that is embedded in the mixed-speech complex. The selective attention-induced unmasking of target speech may be associated with the improved beta-band functional connectivity from the central frontal area to the right motor area, suggesting a top-down attentional modulation of the speech-motor process.</p>
</abstract>
<kwd-group>
<kwd>selective attention</kwd>
<kwd>speech unmasking</kwd>
<kwd>long-term neural activities</kwd>
<kwd>neural network</kwd>
<kwd>motor theory</kwd>
<kwd>informational masking</kwd>
</kwd-group>
<contract-sponsor id="cn001">National Natural Science Foundation of China<named-content content-type="fundref-id">10.13039/501100001809</named-content></contract-sponsor>
<counts>
<fig-count count="5"/>
<table-count count="1"/>
<equation-count count="0"/>
<ref-count count="64"/>
<page-count count="10"/>
<word-count count="0"/>
</counts>
</article-meta>
</front>
<body>
<sec><title>Introduction</title>
<p>The &#x201C;cocktail-party&#x201D; problem (<xref ref-type="bibr" rid="B10">Cherry, 1953</xref>) indicates the astonishing ability of human listeners to recognize target speech in noisy environments with multiple-people talking. It has been confirmed that selective attention plays a critical role in this perceptual/cognitive capacity (e.g., <xref ref-type="bibr" rid="B8">Brungart, 2001</xref>; <xref ref-type="bibr" rid="B21">Freyman et al., 2001</xref>, <xref ref-type="bibr" rid="B22">2004</xref>; <xref ref-type="bibr" rid="B53">Roman et al., 2003</xref>; <xref ref-type="bibr" rid="B38">Li et al., 2004</xref>; <xref ref-type="bibr" rid="B4">Bidet-Caulet et al., 2007</xref>; <xref ref-type="bibr" rid="B19">Ezzatian et al., 2011</xref>; <xref ref-type="bibr" rid="B28">Golumbic et al., 2012</xref>; <xref ref-type="bibr" rid="B44">Mesgarani and Chang, 2012</xref>). On the other hand, non-selective attention provides more generalized and sustain alertness for preparing the emergence of high-priority signals (<xref ref-type="bibr" rid="B48">Posner and Petersen, 1990</xref>, <xref ref-type="bibr" rid="B49">2012</xref>). The relationship between selective attention and non-selective attention has been an attractive issue in the visual research field (e.g., <xref ref-type="bibr" rid="B11">Coull et al., 1998</xref>; <xref ref-type="bibr" rid="B42">Matthias et al., 2010</xref>), but has not been systematically investigated in the auditory research field.</p>
<p>Recently, a few studies on how selective attention affects the cortical representation of target speech have been reported (e.g., <xref ref-type="bibr" rid="B35">Lalor and Foxe, 2010</xref>; <xref ref-type="bibr" rid="B13">Ding and Simon, 2012</xref>, <xref ref-type="bibr" rid="B14">2013</xref>; <xref ref-type="bibr" rid="B27">Golumbic et al., 2013</xref>; <xref ref-type="bibr" rid="B34">Kong et al., 2014</xref>). Particularly, under &#x201C;cocktail-party&#x201D; listening conditions, selective attention modulates low-frequency oscillations of cortical responses to speech stimuli, exhibiting both enhanced tracking of target-speech signals and enhanced suppression of masker-speech signals (<xref ref-type="bibr" rid="B33">Kerlin et al., 2010</xref>; <xref ref-type="bibr" rid="B35">Lalor and Foxe, 2010</xref>; <xref ref-type="bibr" rid="B51">Power et al., 2010</xref>, <xref ref-type="bibr" rid="B50">2012</xref>; <xref ref-type="bibr" rid="B44">Mesgarani and Chang, 2012</xref>; <xref ref-type="bibr" rid="B45">O&#x2019;Sullivan et al., 2014</xref>). It is of interest to know how the neural representation of speech signals under &#x201C;cocktail-party&#x201D; conditions is affected by shifting non-selective attention to selective attention.</p>
<p>It has been proposed that low-frequency (alpha and beta bands) oscillations of cortical activation mainly carries top-down modulation information, while high-frequency (gamma) oscillations mainly carries bottom-up information (<xref ref-type="bibr" rid="B58">Wang, 2010</xref>; <xref ref-type="bibr" rid="B2">Bastos et al., 2012</xref>; <xref ref-type="bibr" rid="B59">Weiss and Mueller, 2012</xref>; <xref ref-type="bibr" rid="B6">Bressler and Richter, 2015</xref>; <xref ref-type="bibr" rid="B25">Friston et al., 2015</xref>; <xref ref-type="bibr" rid="B36">Lewis and Bastiaansen, 2015</xref>). Particularly, top-down signals that come to lower-level brain structure underlies the attentional processing that is associated with the synchrony in the beta frequency band (<xref ref-type="bibr" rid="B29">Hanslmayr et al., 2007</xref>; <xref ref-type="bibr" rid="B62">Womelsdorf and Fries, 2007</xref>; <xref ref-type="bibr" rid="B15">Donner and Siegel, 2011</xref>; <xref ref-type="bibr" rid="B6">Bressler and Richter, 2015</xref>; <xref ref-type="bibr" rid="B54">Saarinen et al., 2015</xref>; <xref ref-type="bibr" rid="B57">Todorovic et al., 2015</xref>). More specifically, for example, beta-band activity is related to various top-down cognitive/perceptual processes (review in <xref ref-type="bibr" rid="B17">Engel and Fries, 2010</xref>), including prediction (<xref ref-type="bibr" rid="B18">Engel et al., 2001</xref>; <xref ref-type="bibr" rid="B1">Ahveninen et al., 2013</xref>; <xref ref-type="bibr" rid="B57">Todorovic et al., 2015</xref>; <xref ref-type="bibr" rid="B37">Lewis et al., 2016</xref>) and motor control (<xref ref-type="bibr" rid="B7">Brittain and Brown, 2014</xref>; <xref ref-type="bibr" rid="B47">Piai et al., 2015</xref>). Also, the beta-band oscillation represents functional connectivity between the frontal cortex and motor cortex in attention tasks (<xref ref-type="bibr" rid="B56">Thorpe et al., 2012</xref>; <xref ref-type="bibr" rid="B47">Piai et al., 2015</xref>). It is of interest to know whether neural oscillations in the beta band are involved in speech unmasking based on selective attention.</p>
<p>The present study investigated whether neural oscillations of scalp-recoded electroencephalogram (EEGs) to multiple-talker (voice) speech are modulated by selective attention and what are the potential underlying mechanisms. EEG signals were recorded from participants who either selectively attended to one of the talker&#x2019;s voice or non-selectively attended to the whole mixed-speech complex. Four frequency bands (theta: 4&#x2013;8 Hz; alpha: 8&#x2013;12 Hz; beta: 13&#x2013;30 Hz; gamma: 30&#x2013;48 Hz) of recorded EEGs were analyzed to reveal both the cortical representation of speech signals and the differences in cortical causal connections between the selective attention condition and the non-selective attention condition. Across EEG correlations were used to estimate whether the cortical speech representation becomes more correlated to the attended target speech under the selective attention condition than the non-selective attention condition.</p>
</sec>
<sec id="s1" sec-type="materials|methods">
<title>Materials and Methods</title>
<sec><title>Participants</title>
<p>Twelve younger adults (five males and seven females) with the mean age of 23.6 years old (from 19 to 25 years old) were recruited from Peking University as the participants in this study. They provided informed consent to participate in this study and were paid a modest stipend for their participation. All the participants were right-handed native Mandarin Chinese speakers with normal and balanced (no more than 15 dB difference between the two ears) pure-tone hearing thresholds between 125 and 8000 Hz. The participants gave their written informed consent for participation in this study. The experimental procedures used in this study were approved by the Committee for Protecting Human and Animal Subjects of the Department of Psychology at Peking University.</p>
</sec>
<sec><title>Speech Stimuli</title>
<p>The speech stimuli used in this study were Chinese &#x201C;nonsense&#x201D; sentences. &#x201C;Nonsense&#x201D; sentences are syntactically correct but not semantically meaningful (e.g., <xref ref-type="bibr" rid="B23">Freyman et al., 1999</xref>; <xref ref-type="bibr" rid="B38">Li et al., 2004</xref>; <xref ref-type="bibr" rid="B64">Yang et al., 2007</xref>; <xref ref-type="bibr" rid="B26">Gao et al., 2014</xref>). Direct English translations of these Chinese sentences are similar but not identical to the English &#x201C;nonsense&#x201D; sentences used in previous studies (<xref ref-type="bibr" rid="B30">Helfer, 1997</xref>; <xref ref-type="bibr" rid="B23">Freyman et al., 1999</xref>, <xref ref-type="bibr" rid="B22">2004</xref>; <xref ref-type="bibr" rid="B38">Li et al., 2004</xref>). For example, the English translation of one Chinese nonsense sentence is &#x201C;That corona removes the crest-span bag&#x201D;. The development of the Chinese &#x201C;nonsense&#x201D; sentences has been described elsewhere (<xref ref-type="bibr" rid="B64">Yang et al., 2007</xref>).</p>
<p>In this study, three different younger-adult female talkers recited the speech stimuli with different sentences. In a typical recording trial, during the mixed-speech presentation when EEGs were recorded (Phase III in <bold>Figure <xref ref-type="fig" rid="F1">1</xref></bold>), the three voices reciting differences sentences were presented at the same time, simulating a &#x201C;cocktail-party&#x201D; listening condition. Before the 3-voice mixed-speech presentation, one of the speech stimuli was presented alone (Phase I in <bold>Figure <xref ref-type="fig" rid="F1">1</xref></bold>) to indicate that either the repeatedly presented speech in the mixed-speech presentation was the target speech when the pre-presented speech was recited by Voice 1 or 2, or there was no particular (single) target speech in the mixed-speech presentation when the pre-presented speech was recited by Voice 3. Consequently, the target speech was determined (when recited by Voice 1 or 2), and the other two speech stimuli formed the masker. In other words, the target speech was presented against a two-talker-speech background. Note that two-talker speech maskers were the most effective in inducing informational masking (<xref ref-type="bibr" rid="B22">Freyman et al., 2004</xref>). Each of the three voices recited different sentences and the sound pressure level of the three voices were the same. The mean duration of the sentences was 3.26 s (ranged from 3.1 to 3.5 s).</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption><p><bold>Illustration of the six phases within each trial of EEG recordings.</bold> Phase I: a trial was started with the presentation of a single-voiced speech (Voice 1, 2, or 3) to indicate which stimulation condition the present trial belonged to [Voice 1, selective attention to Voice 1 (top panel); Voice 2, selective attention to Voice 2 (middle panel); Voice 3, non-selective attention to the whole mixed-speech complex (bottom panel)]. Phase II: a period of silence lasting 1 s. Phase III: the presentation of the mixed three-voiced speech. Phase IV: a period of silence lasting 1.2 s. Phase V and Phase VI: the repetition of Phase III and Phase IV, respectively. Under the selective-attention condition (with Voice 1 or 2), participants were instructed to press a button if they had heard a wrong words probe (yellow waves); under the non-selective-attention condition, participants were instructed to press a button if they heard a click probe (yellow waves). The blue, green, and red waves indicate the single speech of Voice 1, Voice 2, and Voice 3, respectively.</p></caption>
<graphic xlink:href="fnhum-11-00034-g001.tif"/>
</fig>
<p>All speech signals were digitized at a sampling rate of 22.05 kHz using a 24-bit Creative Sound Blaster PCI128 with a built-in anti-aliasing filter (Creative Technology, Ltd., Singapore). All the stimuli, including the single-voice speech, mixed-voice speech, and click sounds were transferred using a Creative Extigy sound blaster and presented to participants at the two ears without any interaural time disparities using two tube-ear inserts (Neuroscan, El Paso, TX, USA). The sound pressure level of a single voice was set at 56 dB SPL, calibrated by a Larson Davis Audiometer Calibration and Electroacoustic Testing System (Audit and System 824, Larson Davis, USA). Since the sound pressure level of the three voices were the same, the signal-to-masker ratio (SMR) was -3 dB when a target speech was determined in the mixed-speech presentation.</p>
</sec>
<sec><title>Electrophysiological Recordings</title>
<p>Scalp EEG recordings (with the reference electrode located on the nose) were conducted in a dim double-walled sound-attenuating booth (EMI Shielded Audiometric Examination Acoustic Suite) that was equipped with a 64-channel NeuroScan SynAmps System (Compumedics Limited, Abbotsford, VIC, Australia). EEG signals were processed with a sample rate of 1000 Hz, on-line amplified 500 times, and low-pass filtered below 200 Hz. Eye movements and eye blinks were recorded from electrodes superior and inferior to the left eye and also at the outer canthi of the two eyes. The impedances of all the recording electrodes were kept below 5 k&#x03A9;.</p>
</sec>
<sec><title>Procedures</title>
<p>The effect of selective attention was estimated by examining the differences in EEGs between the selective attention condition and the non-selective attention condition. Voice 1 and Voice 2 were used as either the target voice or the masking voice, and Voice 3 was used only as the masking voice. There were three stimulation conditions for the mixed-speech presentation: (1) Condition 1: selective attention only to Voice 1, (2) Condition 2: selective attention only to Voice 2, and (3) Condition 3: non-selective attention to the whole speech complex (<bold>Figure <xref ref-type="fig" rid="F1">1</xref></bold>).</p>
<p>In addition to the 3-voice mixed-speech presentation, each of the speech stimuli was presented alone to obtain EEGs to the single-speech presentation.</p>
<p>In this study, five &#x201C;nonsense&#x201D; sentences from a pool with totally 360 sentences were randomly assigned to a participants (two sentences for Voice 1; other two different sentences for Voice 2; one sentence for Voice 3) and different participants listened to difference sentences. For each participants, there were four different mixed-speech presentations.</p>
<p>As shown in <bold>Figure <xref ref-type="fig" rid="F1">1</xref></bold>, each trial contained six phases: In Phase I, a trial was started with the presentation of a single-voiced speech (Voice 1, 2, or 3) with the duration about 3.2 s as the cue to indicate which stimulation condition the present trial belonged to (Voice 1, Condition 1; Voice 2, Condition 2; Voice 3, Condition 3). Phase I was followed by Phase II, which was a period of silence lasting 1 s.</p>
<p>In Phase III, the mixed three-voiced speech (about 3.2 s) was presented (the same stimuli under different conditions for a participant). Phase IV was also a period of silence lasting 1.2 s. The Phase V and Phase VI were the repetition of Phase III and Phase IV, respectively. In other words, the mixed speech presentation occurred twice in a trial.</p>
<p>Under a selective attention condition (Condition 1 or 2), participants were instructed to pay attention to the target voice and press a button if they had heard a novel &#x201C;predicate-object&#x201D; structure presented with the same voice as the to-be-attended talker (as the false-word probe, with four syllables and the possibility of 14.2%, <bold>Figure <xref ref-type="fig" rid="F1">1</xref></bold>). Under non-selective attention condition (Condition 3), the participants were instructed to pay attention to the whole speech complex and press a button if they heard a &#x201C;click&#x201D; (as the probe with the possibility of 14.2%) at a random time position (<bold>Figure <xref ref-type="fig" rid="F1">1</xref></bold>). To ensure that the participants could understand and follow the instructions, a training session was conducted before EEG recordings. The percent correct in detecting the probe in each of the participants were required to be no less than 85%.</p>
<p>In total, there were 96 stimulation presentations for EEG recordings (after the removal of the presentations with probes) for each of the three conditions, and these 96 presentations were randomly assigned into four blocks. Each block contained 24 stimulation presentations for each of the three conditions whose presenting order was arranged randomly for a participant. It took about 10 mins to complete one block. To limit eye movements, participants were also asked to stare a cross in the front in a trial.</p>
</sec>
<sec><title>Data Analyses</title>
<p>Using the EEGLAB toolbox (<xref ref-type="bibr" rid="B12">Delorme and Makeig, 2004</xref>) in MATLAB, raw EEG data were filtered by three different band-pass filters (alpha: 8&#x2013;12 Hz; beta: 12&#x2013;30 Hz; gamma: 30&#x2013;48 Hz), and then segmented into epochs from -300 to 3500 ms relative to the onset of a mixed-speech presentation. The baseline correction was conducted in the period of -300 to 0 ms before the presentation onset. The epochs that contained more than &#x00B1;30 &#x03BC;V potential were rejected as artifacts. The rest of epochs were averaged for each condition to analyze the grange causality and across EEG correlations.</p>
<p>To avoid the onset and offset (above 3000-ms) effect (<xref ref-type="bibr" rid="B46">Pasley et al., 2012</xref>), the period of interest was defined within the time 800&#x2013;2800 ms after the mixed-speech presentation onset. The across EEG correlation was calculated by the <italic>corr</italic> function in MATLAB. The Granger Causality (GC) analysis was calculated using the Brainstorm toolbox (<xref ref-type="bibr" rid="B55">Tadel et al., 2011</xref>)<sup><xref ref-type="fn" rid="fn01">1</xref></sup> in the MATLAB environment to estimate causal connectivity associated with the selective attention effect.</p>
<p>Six areas were defined for GC analyses: (1) the left frontal area, including sites F5, F3, F1, FC5, FC3, FC1; (2) the central frontal area including sites, including F3, F1, F2, FC3, FC1, FC2; (3) the right frontal area, including sites F6, F4, F2, FC6, FC4, FC2; (4) the left motor area, including sites C5, C3, C1, CP5, CP3, CP1; (5) the central motor area, including sites C3, C1, C2, CP3, CP1, CP2; (6) the right motor area, including sites C6, C4, C2, CP6, CP4, CP2. The areal GC value for each participant was averaged by the GC values of all site connections in each area.</p>
<p>The change index was calculated as: (v1 - v2)/(v1 + v2), where the v1 and v2 were the value under two different conditions.</p>
<p>Statistical analyses were performed with IBM SPSS Statistics 20 (SPSS Inc., Chicago, IL, USA). Within-participants, paired <italic>t</italic>-tests and Pearson correlation were conducted to assess differences between conditions. The null-hypothesis rejection level was set at 0.05.</p>
</sec>
</sec>
<sec><title>Results</title>
<sec><title>The Effect of Selective Attention on Cortical Representations of Speech Signals against Speech Masking</title>
<p>To estimate the effect of selective attention on cortical representation of speech against speech masking, Pearson correlation coefficients were calculated between the EEGs to the mixed-speech complex under the selective attention condition (when only one voice was attended) and the EEGs to a single-voiced speech, which was used as either the attended one (the target voice) or not the attended one in the mixed-speech complex.</p>
<p>As showed in <bold>Figure <xref ref-type="fig" rid="F2">2</xref></bold>, the 5-Hz high-pass filtered all-site-averaged EEGs to the mixed-speech complex were significantly more correlated to the 5-Hz high-pass filtered all-site-averaged ERPs to the single speech that was used as the target speech in the speech complex than the EEGs to the single speech that was not attended in the speech complex [<italic>t</italic>(11) = 3.124, <italic>p</italic> = 0.010, paired <italic>t</italic>-test]. These results suggested that selective attention significantly improved the cortical representation of target-speech signals in a multi-talker environment (Supplementary Figure <xref ref-type="supplementary-material" rid="SM1">S1</xref>).</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption><p><bold>Under the selective attention condition, the correlation between the all-site-averaged EEGs to the mixed-speech complex and the all-site-averaged EEGs to the single speech that was either the target or the masker speech in the mixed-speech complex.</bold> <sup>&#x2217;&#x2217;</sup><italic>p</italic> &#x003C; 0.01, paired <italic>t</italic>-test. The error bar indicates the standard errors of the mean.</p></caption>
<graphic xlink:href="fnhum-11-00034-g002.tif"/>
</fig>
<p>To further estimate whether different frequency-band oscillations in EEGs were differently affected by selective attention, EEG data for each of the various frequency bands (theta, alpha, beta, gamma, and broad) were analyzed separately. In <bold>Figure <xref ref-type="fig" rid="F3">3</xref></bold>, for each of the frequency bands, the first left column shows the absolute correlation coefficients between the EEGs to the mixed-speech complex under the non-selective attention condition (NS) and the EEGs to a single speech for all the recording sites; the second left column shows the absolute correlation coefficients between the EEGs to the mixed speech under the selective attention condition (S) and the EEGs to the target single speech for all the recording sites.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption><p><bold>The two left columns: for each of the five types of frequency bands [theta (&#x1D703;), alpha (&#x03B1;), beta (&#x03B2;), gamma (&#x03B3;), broad], the scalp topographical maps showing location distributions of absolute correlations between the EEGs to the mixed-speech complex and EEGs to a single speech under either the non-selective attention (NS) condition or the selective attention (S) condition.</bold> The two right columns: for each of the frequency bands, the recordings sites at which the correlation difference between the two attention conditions was significant when the <italic>p</italic> level was either 0.05 and or 0.015.</p></caption>
<graphic xlink:href="fnhum-11-00034-g003.tif"/>
</fig>
<p>To reveal the frequency band that was the most vulnerable to selective attention, <bold>Figure <xref ref-type="fig" rid="F3">3</xref></bold> also shows the statistically thresholded topographical map (the two right columns) indicating the electrode sites exhibiting significant differences in absolute correlation coefficient between the selective attention condition (S) and the non-selective attention condition (NS). When the <italic>p</italic> level was 0.05 (the second right column), both beta- and gamma-band components of EEGs recorded from a few electrode sites exhibited significant differences between the two attention conditions. <bold>Table <xref ref-type="table" rid="T1">1</xref></bold> shows the <italic>p</italic>-values for these electrode sites. Also shown in <bold>Table <xref ref-type="table" rid="T1">1</xref></bold>, only the beta-band component of EEGs recorded from the site Cz exhibited a significant difference between the two attentional conditions when the <italic>p</italic> was as low as 0.011. In other words, the beta-band obtained at the site Cz was the only component exhibiting a significant difference between the two attention conditions when the <italic>p</italic> value was less than 0.020. The right column in <bold>Figure <xref ref-type="fig" rid="F3">3</xref></bold> presents the results indicating that the beta-band component of EEGs at the site Cz was the only one exhibiting a significant difference between the two attention conditions when the <italic>p</italic> value was 0.015 (which was just larger than 0.011 but smaller than 0.020). More in detail, at the <italic>p</italic> level of 0.015, the mixed-speech-evoked EEGs at site Cz were significantly more correlated with the single-speech-evoked EEGs under the selective attention condition than under the non-selective attention condition for beta band [<italic>t</italic>(11) = 3.029, <italic>p</italic> = 0.011, paired <italic>t</italic>-test], but not for other bands (both <italic>p</italic> > 0.05, paired <italic>t</italic>-test), indicating that the EEG beta-band component at the site Cz was the most vulnerable to selective attention (Supplementary Figure <xref ref-type="supplementary-material" rid="SM2">S2</xref>).</p>
<table-wrap position="float" id="T1">
<label>Table 1</label>
<caption><p>Electrode sites at which beta and gamma bands were significantly different between the two attention conditions.</p></caption>
<table cellspacing="5" cellpadding="5" frame="hsides" rules="groups">
<thead>
<tr>
<th valign="top" align="left">Band</th>
<th valign="top" align="center">Sites</th>
<th valign="top" align="center"><italic>df</italic></th>
<th valign="top" align="center"><italic>t</italic></th>
<th valign="top" align="center"><italic>p</italic></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">Beta</td>
<td valign="top" align="center">CZ</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">3.029</td>
<td valign="top" align="center">0.011</td></tr>
<tr>
<td valign="top" align="left">Beta</td>
<td valign="top" align="center">F7</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.642</td>
<td valign="top" align="center">0.023</td>
</tr>
<tr>
<td valign="top" align="left">Beta</td>
<td valign="top" align="center">F1</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.494</td>
<td valign="top" align="center">0.030</td></tr>
<tr>
<td valign="top" align="left">Beta</td>
<td valign="top" align="center">FT7</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.408</td>
<td valign="top" align="center">0.035</td>
</tr>
<tr>
<td valign="top" align="left">Beta</td>
<td valign="top" align="center">F3</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.384</td>
<td valign="top" align="center">0.036</td></tr>
<tr>
<td valign="top" align="left">Beta</td>
<td valign="top" align="center">FP1</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.325</td>
<td valign="top" align="center">0.040</td>
</tr>
<tr>
<td valign="top" align="left">Beta</td>
<td valign="top" align="center">FPZ</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.258</td>
<td valign="top" align="center">0.045</td></tr>
<tr>
<td valign="top" align="left">Beta</td>
<td valign="top" align="center">F5</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.203</td>
<td valign="top" align="center">0.050</td>
</tr>
<tr>
<td valign="top" align="left">Gamma</td>
<td valign="top" align="center">PO7</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.593</td>
<td valign="top" align="center">0.025</td></tr>
<tr>
<td valign="top" align="left">Gamma</td>
<td valign="top" align="center">PO5</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.554</td>
<td valign="top" align="center">0.027</td>
</tr>
<tr>
<td valign="top" align="left">Gamma</td>
<td valign="top" align="center">P5</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.372</td>
<td valign="top" align="center">0.037</td></tr>
<tr>
<td valign="top" align="left">Gamma</td>
<td valign="top" align="center">TP7</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.311</td>
<td valign="top" align="center">0.041</td>
</tr>
<tr>
<td valign="top" align="left">Gamma</td>
<td valign="top" align="center">PO3</td>
<td valign="top" align="center">11</td>
<td valign="top" align="center">2.283</td>
<td valign="top" align="center">0.043</td></tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec><title>Beta-Band Causal Connectivity Enhanced by Selective Attention</title>
<p>Since the beta-band component in EEGs to speech was significantly enhanced by selective attention, it is of importance to know whether some causal connectivities (i.e., GCs) of beta band were also enhanced by selective attention. The results of GC analyses showed that the following four beta-band GCs were significantly facilitated by selective attention (<italic>p</italic> &#x003C; 0.05, paired <italic>t</italic>-test, <bold>Figure <xref ref-type="fig" rid="F4">4</xref></bold>), including the ones (1) from site FT7 to the right motor area [Voice 1, <italic>t</italic>(11) = 2.769, <italic>p</italic> = 0.018; Voice 2, <italic>t</italic>(11) = 2.371, <italic>p</italic> = 0.037], (2) from the left frontal area to the right motor area [Voice 1, <italic>t</italic>(11) = 3.223, <italic>p</italic> = 0.008; Voice 2, <italic>t</italic>(11) = 2.629, <italic>p</italic> = 0.023], (3) from the central frontal area to the right motor area [Voice 1, <italic>t</italic>(11) = 2.344, <italic>p</italic> = 0.039; Voice 2, <italic>t</italic>(11) = 2.451, <italic>p</italic> = 0.032], and (4) from the central frontal area to the right frontal area [Voice 1, <italic>t</italic>(11) = 3.895, <italic>p</italic> = 0.002; Voice 2, <italic>t</italic>(11) = 2.692, <italic>p</italic> = 0.021].</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption><p><bold>The four significant event-related Granger Causalities (GCs) induced by selective attention (S) against non-selective attention (NS) of beta band (<italic>p</italic> &#x003C; 0.05)</bold>.</p></caption>
<graphic xlink:href="fnhum-11-00034-g004.tif"/>
</fig>
</sec>
<sec><title>Correlation between Causal Connectivity and Cortical Representation of Speech against Speech Masking</title>
<p>As described above, selective attention enhanced both the beta-band component of the cortical representation of the target speech in mixed-speech complex and the four GCs (the ones from site FT7 to the right motor area, from the left frontal area to the right motor area, from the central frontal area to the right motor area, from the central frontal area to the right frontal area). Thus, it is of interest to know whether the selective attention-induced beta-band improvement of the speech representation (measured by the correlation change index, see below) was correlated with the selective-attention-induced improvement of any of the four GCs (measured by the GC change index, see below).</p>
<p>The correlation change index induced by selective attention was calculated as: (&#x03C1;<sub>S</sub> <sub>-</sub> &#x03C1;<sub>NS</sub>)/(&#x03C1;<sub>S</sub> + &#x03C1;<sub>NS</sub>), where &#x03C1;<sub>S</sub> and &#x03C1;<sub>NS</sub> were the beta-band EEG correlation coefficients between mixed-speech stimulation and single-speech stimulation at site Cz under the selective attention condition (S) and under the non-selective attention condition (NS), respectively. The positive value represented a selective attentional predominance while the negative value represented a non-selective attentional predominance.</p>
<p>The GC change index for a frequency band (such as beta band) was calculated as: (G<sub>S,c</sub> <sub>-</sub> G<sub>NS,c</sub>)/(G<sub>S,c</sub> + G<sub>NS,c</sub>), where G<sub>S</sub> and G<sub>NS</sub> were the Granger Causalities for a connection <italic>c</italic> under the selective (S) attention condition and the non-selective (NS) attention condition, respectively.</p>
<p>The results showed that the correlation change index of beta band across participants was significantly correlated with the beta-band GC change index only for connectivity from the central frontal area to the right motor area (<italic>r</italic> = 0.585, <italic>p</italic> = 0.046; <bold>Figure <xref ref-type="fig" rid="F5">5</xref></bold>).</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption><p><bold>For each of the four significant GCs shown in <bold>Figure <xref ref-type="fig" rid="F4">4</xref></bold>, the correlation between the beta (&#x03B2;)-band correlation change index induced by selective attention and the beta (&#x03B2;)-band GC change index induced by selective attention. (A)</bold> Causal connectivity from the central frontal area to the right motor area; <bold>(B)</bold> Causal connectivity from the left frontal area to the right motor area; <bold>(C)</bold> Causal connectivity from the central frontal area to the right frontal area; <bold>(D)</bold> Causal connectivity from site TP7 to the right motor area. <sup>&#x2217;</sup><italic>p</italic> &#x003C; 0.05.</p></caption>
<graphic xlink:href="fnhum-11-00034-g005.tif"/>
</fig>
</sec>
</sec>
<sec><title>Discussion</title>
<p>By recording scalp EEGs to speech stimuli, this study investigated under a simulated &#x201C;cocktail&#x201D; party condition with speech-on-speech masking, how selective attention modulates cortical representation of the masked target speech. Note that these findings on the difference between selective-attention and non-selective-attention conditions are based on the use of speech sounds. It is of importance to know whether similar findings can be obtained using non-speech sounds.</p>
<sec><title>Selective Attention Improves the Cortical Representation of Target-Speech Signals</title>
<p>Previous studies have reported that human cortical oscillations represent temporal structures of speech signals with high fidelity (<xref ref-type="bibr" rid="B13">Ding and Simon, 2012</xref>; <xref ref-type="bibr" rid="B44">Mesgarani and Chang, 2012</xref>). The results of this study showed that the correlation between the all-site-averaged EEGs to the mixed-speech complex and the all-site-averaged EEGs to the single speech that was used as the target in the speech complex was significantly larger than the correlation between the all-site-averaged EEGs to the mixed-speech complex and the all-site-averaged EEGs to the single speech that was not attended in the speech complex. Thus, this study supports the view that under a speech-on-speech masking condition, selective attention to a single-voice speech improves the cortical representation of this target single-voice speech (<xref ref-type="bibr" rid="B13">Ding and Simon, 2012</xref>, <xref ref-type="bibr" rid="B14">2013</xref>; <xref ref-type="bibr" rid="B44">Mesgarani and Chang, 2012</xref>; <xref ref-type="bibr" rid="B27">Golumbic et al., 2013</xref>; <xref ref-type="bibr" rid="B45">O&#x2019;Sullivan et al., 2014</xref>).</p>
</sec>
<sec><title>The Beta-Band Component of the EEGs to Speech Is the Most Vulnerable to Selective Attention</title>
<p>In this study, following EEG data for each of the three frequency bands (alpha, beta, and gamma) were analyzed separately, the results showed that the beta-band component, but not either the alpha-band component or the gamma-band component, in the mixed-speech-evoked EEGs, was significantly more correlated with the single-speech-evoked EEGs under the selective attention condition (where the target single-voice speech was attended) than under the non-selective attention condition. Thus, the EEG beta-band component was the most vulnerable to selective attention.</p>
<p>Beta oscillations are associated with attention and predictions (<xref ref-type="bibr" rid="B17">Engel and Fries, 2010</xref>; <xref ref-type="bibr" rid="B15">Donner and Siegel, 2011</xref>; <xref ref-type="bibr" rid="B59">Weiss and Mueller, 2012</xref>; <xref ref-type="bibr" rid="B57">Todorovic et al., 2015</xref>), which are critical to speech cognition. Particularly, the top-down propagation of predictions reflected by beta oscillations (<xref ref-type="bibr" rid="B18">Engel et al., 2001</xref>; <xref ref-type="bibr" rid="B2">Bastos et al., 2012</xref>; <xref ref-type="bibr" rid="B1">Ahveninen et al., 2013</xref>; <xref ref-type="bibr" rid="B36">Lewis and Bastiaansen, 2015</xref>; <xref ref-type="bibr" rid="B57">Todorovic et al., 2015</xref>; <xref ref-type="bibr" rid="B37">Lewis et al., 2016</xref>) may be more critical for selective-attention-induced unmasking of speech, probably through enhancing the mechanism underlying binding distributed sets of neurons into a coherent representation of speech contents (<xref ref-type="bibr" rid="B59">Weiss and Mueller, 2012</xref>).</p>
</sec>
<sec><title>Selective-Attention Facilitated Beta-Band Causal Connectivity from the Central Frontal Area to the Right Motor Area</title>
<p>The results of this study also showed that in total four beta-band causal connectivities (measured as GCs) were enhanced by selective attention, including the ones (1) from site FT7 to the right motor area, (2) from the left frontal area to the right motor area, (3) from the central frontal area to the right motor area, and (4) from the central frontal area to the right frontal area. However, only the selective-attention-induced enhancement of beta-band GC from the central frontal area to the right motor area was significantly correlated to the selective-attention-induced enhancement of the correlation between beta-band oscillations to the mixed speech complex and beta-band oscillations to the single speech. The results suggest that the selective-attention-induced improvement of beta-band representation of target speech signals is associated with the enhanced top-down modulation of the motor areas in the right hemisphere by the central frontal cortical areas. In other words, selective attention improves speech-related motor processes. However, due to the low spatial resolution of EEGs, whether the beta activities over central areas are based on the auditory or motor activity need further investigation in the future.</p>
<p>The <italic>Motor Theory</italic> of speech perception proposes that the interaction between the auditory and motor systems plays an essential role in speech perception (<xref ref-type="bibr" rid="B40">Liberman et al., 1952</xref>, <xref ref-type="bibr" rid="B39">1967</xref>; <xref ref-type="bibr" rid="B41">Liberman and Mattingly, 1985</xref>; for review see <xref ref-type="bibr" rid="B63">Wu et al., 2014</xref>). It has been evident that speech perception activates the motor cortex (<xref ref-type="bibr" rid="B20">Fadiga et al., 2002</xref>; <xref ref-type="bibr" rid="B9">Callan et al., 2004</xref>; <xref ref-type="bibr" rid="B61">Wilson et al., 2004</xref>; <xref ref-type="bibr" rid="B52">Pulverm&#x00FC;ller et al., 2006</xref>; <xref ref-type="bibr" rid="B60">Wilson and Iacoboni, 2006</xref>; <xref ref-type="bibr" rid="B43">Meister et al., 2007</xref>; <xref ref-type="bibr" rid="B3">Bever and Poeppel, 2010</xref>; <xref ref-type="bibr" rid="B31">Hickok et al., 2011</xref>; <xref ref-type="bibr" rid="B16">Elemans et al., 2015</xref>). Thus, under adverse listening conditions (such as the cocktail-party environment) where the perceptual load is high (<xref ref-type="bibr" rid="B32">Hickok and Poeppel, 2007</xref>; <xref ref-type="bibr" rid="B24">Fridriksson et al., 2008</xref>; <xref ref-type="bibr" rid="B5">Bishop and Miller, 2009</xref>), with the involvement of the motor system the listener can better identify the speaker&#x2019;s intention and follow the target stream (<xref ref-type="bibr" rid="B63">Wu et al., 2014</xref>).</p>
</sec>
</sec>
<sec><title>Conclusion</title>
<list list-type="simple" prefix-word="simple">
<list-item><label>(1)</label><p>The cortical representation of target-speech signals under the multiple-people talking condition is specifically improved by selective attention, and the beta-band EEG component is the most vulnerable to selective attention.</p></list-item>
<list-item><label>(2)</label><p>The selective-attention-induced enhancement of beta-band causal connectivity from the central frontal area to the right motor area is correlated with the selective-attention-induced enhancement of the cortical beta-band representation of target speech.</p></list-item>
<list-item><label>(3)</label><p>Selective attention to a single-voiced target speech, which is embedded in a mixed-speech complex (with speech-on-speech masking), improves the cortical representation of the target speech by facilitating the top-down frontal modulation of the motor cortical areas.</p></list-item>
<list-item><label>(4)</label><p>The unmasking of target speech based on selective attention may be caused by top-down attentional modulation of the speech-motor interactions.</p></list-item>
</list>
</sec>
<sec><title>Author Contributions</title>
<p>YG, QW, and YD: Experimental design, experiment set up, experiment conduction, data analyses, figure/table construction, and paper writing. CW and HL: Experimental design, data analyses, and paper writing. XW: Experimental design and paper writing. LL: Experimental design, figure/table construction, and paper writing. TQ: Experimental design, experiment set up, and paper writing.</p>
</sec>
<sec><title>Conflict of Interest Statement</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
</body>
<back>
<ack>
<p>This work was supported by supported by the &#x2018;973&#x2019; National Basic Research Program of China (2015CB351800), the National High Technology Research and Development Program of China (863 Program: 2015AA016306), the Beijing Municipal Science and Tech Commission (Z161100002616017), and the National Natural Science Foundation of China (81501155, 61171186, 61671187).</p>
</ack>
<sec sec-type="supplementary material">
<title>Supplementary Material</title>
<p>The Supplementary Material for this article can be found online at: <ext-link ext-link-type="uri" xlink:href="http://journal.frontiersin.org/article/10.3389/fnhum.2017.00034/full#supplementary-materia">http://journal.frontiersin.org/article/10.3389/fnhum.2017.00034/full#supplementary-materia</ext-link></p>
<supplementary-material xlink:href="Image_1.pdf" id="SM1" mimetype="application/pdf" xmlns:xlink="http://www.w3.org/1999/xlink"/>
<supplementary-material xlink:href="Image_2.pdf" id="SM2" mimetype="application/pdf" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<ref-list>
<title>References</title>
<ref id="B1"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ahveninen</surname> <given-names>J.</given-names></name> <name><surname>Huang</surname> <given-names>S.</given-names></name> <name><surname>Belliveau</surname> <given-names>J. W.</given-names></name> <name><surname>Chang</surname> <given-names>W. T.</given-names></name> <name><surname>H&#x00E4;m&#x00E4;l&#x00E4;inen</surname> <given-names>M.</given-names></name></person-group> (<year>2013</year>). <article-title>Dynamic oscillatory processes governing cued orienting and allocation of auditory attention.</article-title> <source><italic>J. Cogn. Neurosci.</italic></source> <volume>25</volume> <fpage>1926</fpage>&#x2013;<lpage>1943</lpage>. <pub-id pub-id-type="doi">10.1162/jocn_a_00452</pub-id></citation></ref>
<ref id="B2"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bastos</surname> <given-names>A. M.</given-names></name> <name><surname>Usrey</surname> <given-names>W. M.</given-names></name> <name><surname>Adams</surname> <given-names>R. A.</given-names></name> <name><surname>Mangun</surname> <given-names>G. R.</given-names></name> <name><surname>Fries</surname> <given-names>P.</given-names></name> <name><surname>Friston</surname> <given-names>K. J.</given-names></name></person-group> (<year>2012</year>). <article-title>Canonical microcircuits for predictive coding.</article-title> <source><italic>Neuron</italic></source> <volume>76</volume> <fpage>695</fpage>&#x2013;<lpage>711</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuron.2012.10.038</pub-id></citation></ref>
<ref id="B3"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bever</surname> <given-names>T. G.</given-names></name> <name><surname>Poeppel</surname> <given-names>D.</given-names></name></person-group> (<year>2010</year>). <article-title>Analysis by synthesis: a (re-) emerging program of research for language and vision.</article-title> <source><italic>Biolinguistics</italic></source> <volume>4</volume> <fpage>174</fpage>&#x2013;<lpage>200</lpage>.</citation></ref>
<ref id="B4"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bidet-Caulet</surname> <given-names>A.</given-names></name> <name><surname>Fischer</surname> <given-names>C.</given-names></name> <name><surname>Besle</surname> <given-names>J.</given-names></name> <name><surname>Aguera</surname> <given-names>P. E.</given-names></name> <name><surname>Giard</surname> <given-names>M. H.</given-names></name> <name><surname>Bertrand</surname> <given-names>O.</given-names></name></person-group> (<year>2007</year>). <article-title>Effects of selective attention on the electrophysiological representation of concurrent sounds in the human auditory cortex.</article-title> <source><italic>J. Neurosci.</italic></source> <volume>27</volume> <fpage>9252</fpage>&#x2013;<lpage>9261</lpage>. <pub-id pub-id-type="doi">10.1523/JNEUROSCI.1402-07.2007</pub-id></citation></ref>
<ref id="B5"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bishop</surname> <given-names>C. W.</given-names></name> <name><surname>Miller</surname> <given-names>L. M.</given-names></name></person-group> (<year>2009</year>). <article-title>A multisensory cortical network for understanding speech in noise.</article-title> <source><italic>J. Cogn. Neurosci.</italic></source> <volume>21</volume> <fpage>1790</fpage>&#x2013;<lpage>1804</lpage>. <pub-id pub-id-type="doi">10.1162/jocn.2009.21118</pub-id></citation></ref>
<ref id="B6"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bressler</surname> <given-names>S. L.</given-names></name> <name><surname>Richter</surname> <given-names>C. G.</given-names></name></person-group> (<year>2015</year>). <article-title>Interareal oscillatory synchronization in top-down neocortical processing.</article-title> <source><italic>Curr. Opin. Neurobiol.</italic></source> <volume>31</volume> <fpage>62</fpage>&#x2013;<lpage>66</lpage>. <pub-id pub-id-type="doi">10.1016/j.conb.2014.08.010</pub-id></citation></ref>
<ref id="B7"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Brittain</surname> <given-names>J. S.</given-names></name> <name><surname>Brown</surname> <given-names>P.</given-names></name></person-group> (<year>2014</year>). <article-title>Oscillations and the basal ganglia: motor control and beyond.</article-title> <source><italic>Neuroimage</italic></source> <volume>85</volume> <fpage>637</fpage>&#x2013;<lpage>647</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuroimage.2013.05.084</pub-id></citation></ref>
<ref id="B8"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Brungart</surname> <given-names>D. S.</given-names></name></person-group> (<year>2001</year>). <article-title>Informational and energetic masking effects in the perception of two simultaneous talkers.</article-title> <source><italic>J. Acoust. Soc. Am.</italic></source> <volume>109</volume> <fpage>1101</fpage>&#x2013;<lpage>1109</lpage>. <pub-id pub-id-type="doi">10.1121/1.1345696</pub-id></citation></ref>
<ref id="B9"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Callan</surname> <given-names>D. E.</given-names></name> <name><surname>Jones</surname> <given-names>J. A.</given-names></name> <name><surname>Callan</surname> <given-names>A. M.</given-names></name> <name><surname>Akahane-Yamada</surname> <given-names>R.</given-names></name></person-group> (<year>2004</year>). <article-title>Phonetic perceptual identification by native-and second-language speakers differentially activates brain regions involved with acoustic phonetic processing and those involved with articulatory&#x2013;auditory/orosensory internal models.</article-title> <source><italic>Neuroimage</italic></source> <volume>22</volume> <fpage>1182</fpage>&#x2013;<lpage>1194</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuroimage.2004.03.006</pub-id></citation></ref>
<ref id="B10"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cherry</surname> <given-names>E. C.</given-names></name></person-group> (<year>1953</year>). <article-title>Some experiments on the recognition of speech, with one and with two ears.</article-title> <source><italic>J. Acoust. Soc. Am.</italic></source> <volume>25</volume> <fpage>975</fpage>&#x2013;<lpage>979</lpage>. <pub-id pub-id-type="doi">10.1121/1.1907229</pub-id></citation></ref>
<ref id="B11"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Coull</surname> <given-names>J. T.</given-names></name> <name><surname>Frackowiak</surname> <given-names>R. S. J.</given-names></name> <name><surname>Frith</surname> <given-names>C. D.</given-names></name></person-group> (<year>1998</year>). <article-title>Monitoring for target objects: activation of right frontal and parietal cortices with increasing time on task.</article-title> <source><italic>Neuropsychologia</italic></source> <volume>36</volume> <fpage>1325</fpage>&#x2013;<lpage>1334</lpage>. <pub-id pub-id-type="doi">10.1016/S0028-3932(98)00035-9</pub-id></citation></ref>
<ref id="B12"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Delorme</surname> <given-names>A.</given-names></name> <name><surname>Makeig</surname> <given-names>S.</given-names></name></person-group> (<year>2004</year>). <article-title>EEGLAB: an open source toolbox for analysis of single-trial EEG dynamics including independent component analysis.</article-title> <source><italic>J. Neurosci. Methods</italic></source> <volume>134</volume> <fpage>9</fpage>&#x2013;<lpage>21</lpage>. <pub-id pub-id-type="doi">10.1016/j.jneumeth.2003.10.009</pub-id></citation></ref>
<ref id="B13"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ding</surname> <given-names>N.</given-names></name> <name><surname>Simon</surname> <given-names>J. Z.</given-names></name></person-group> (<year>2012</year>). <article-title>Emergence of neural encoding of auditory objects while listening to competing speakers.</article-title> <source><italic>Proc. Natl. Acad. Sci. U.S.A.</italic></source> <volume>109</volume> <fpage>11854</fpage>&#x2013;<lpage>11859</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.1205381109</pub-id></citation></ref>
<ref id="B14"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ding</surname> <given-names>N.</given-names></name> <name><surname>Simon</surname> <given-names>J. Z.</given-names></name></person-group> (<year>2013</year>). <article-title>Adaptive temporal encoding leads to a background-insensitive cortical representation of speech.</article-title> <source><italic>J. Neurosci.</italic></source> <volume>33</volume> <fpage>5728</fpage>&#x2013;<lpage>5735</lpage>. <pub-id pub-id-type="doi">10.1523/JNEUROSCI.5297-12.2013</pub-id></citation></ref>
<ref id="B15"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Donner</surname> <given-names>T. H.</given-names></name> <name><surname>Siegel</surname> <given-names>M.</given-names></name></person-group> (<year>2011</year>). <article-title>A framework for local cortical oscillation patterns.</article-title> <source><italic>Trends Cogn. Sci.</italic></source> <volume>15</volume> <fpage>191</fpage>&#x2013;<lpage>199</lpage>. <pub-id pub-id-type="doi">10.1016/j.tics.2011.03.007</pub-id></citation></ref>
<ref id="B16"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Elemans</surname> <given-names>C. P. H.</given-names></name> <name><surname>Rasmussen</surname> <given-names>J. H.</given-names></name> <name><surname>Herbst</surname> <given-names>C. T.</given-names></name> <name><surname>D&#x00FC;ring</surname> <given-names>D. N.</given-names></name> <name><surname>Zollinger</surname> <given-names>S. A.</given-names></name> <name><surname>Brumm</surname> <given-names>H.</given-names></name><etal/></person-group> (<year>2015</year>). <article-title>Universal mechanisms of sound production and control in birds and mammals.</article-title> <source><italic>Nat. Commun.</italic></source> <volume>6</volume>:<issue>8978</issue>. <pub-id pub-id-type="doi">10.1038/ncomms9978</pub-id></citation></ref>
<ref id="B17"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Engel</surname> <given-names>A. K.</given-names></name> <name><surname>Fries</surname> <given-names>P.</given-names></name></person-group> (<year>2010</year>). <article-title>Beta-band oscillations&#x2014;signalling the status quo?</article-title> <source><italic>Curr. Opin. Neurobiol.</italic></source> <volume>20</volume> <fpage>156</fpage>&#x2013;<lpage>165</lpage>. <pub-id pub-id-type="doi">10.1016/j.conb.2010.02.015</pub-id></citation></ref>
<ref id="B18"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Engel</surname> <given-names>A. K.</given-names></name> <name><surname>Fries</surname> <given-names>P.</given-names></name> <name><surname>Singer</surname> <given-names>W.</given-names></name></person-group> (<year>2001</year>). <article-title>Dynamic predictions: oscillations and synchrony in top&#x2013;down processing.</article-title> <source><italic>Nat. Rev. Neurosci.</italic></source> <volume>2</volume> <fpage>704</fpage>&#x2013;<lpage>716</lpage>. <pub-id pub-id-type="doi">10.1038/35094565</pub-id></citation></ref>
<ref id="B19"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ezzatian</surname> <given-names>P.</given-names></name> <name><surname>Li</surname> <given-names>L. A.</given-names></name> <name><surname>Pichora-Fuller</surname> <given-names>K.</given-names></name> <name><surname>Schneider</surname> <given-names>B.</given-names></name></person-group> (<year>2011</year>). <article-title>The effect of priming on release from informational masking is equivalent for younger and older adults.</article-title> <source><italic>Ear Hear.</italic></source> <volume>32</volume> <fpage>84</fpage>&#x2013;<lpage>96</lpage>. <pub-id pub-id-type="doi">10.1097/AUD.0b013e3181ee6b8a</pub-id></citation></ref>
<ref id="B20"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fadiga</surname> <given-names>L.</given-names></name> <name><surname>Craighero</surname> <given-names>L.</given-names></name> <name><surname>Buccino</surname> <given-names>G.</given-names></name> <name><surname>Rizzolatti</surname> <given-names>G.</given-names></name></person-group> (<year>2002</year>). <article-title>Speech listening specifically modulates the excitability of tongue muscles: a TMS study.</article-title> <source><italic>Eur. J. Neurosci.</italic></source> <volume>15</volume> <fpage>399</fpage>&#x2013;<lpage>402</lpage>. <pub-id pub-id-type="doi">10.1046/j.0953-816x.2001.01874.x</pub-id></citation></ref>
<ref id="B21"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Freyman</surname> <given-names>R. L.</given-names></name> <name><surname>Balakrishnan</surname> <given-names>U.</given-names></name> <name><surname>Helfer</surname> <given-names>K. S.</given-names></name></person-group> (<year>2001</year>). <article-title>Spatial release from informational masking in speech recognition.</article-title> <source><italic>J. Acoust. Soc. Am.</italic></source> <volume>109</volume> <fpage>2112</fpage>&#x2013;<lpage>2122</lpage>. <pub-id pub-id-type="doi">10.1121/1.1354984</pub-id></citation></ref>
<ref id="B22"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Freyman</surname> <given-names>R. L.</given-names></name> <name><surname>Balakrishnan</surname> <given-names>U.</given-names></name> <name><surname>Helfer</surname> <given-names>K. S.</given-names></name></person-group> (<year>2004</year>). <article-title>Effect of number of masking talkers and auditory priming on informational masking in speech recognition.</article-title> <source><italic>J. Acoust. Soc. Am.</italic></source> <volume>115</volume> <fpage>2246</fpage>&#x2013;<lpage>2256</lpage>. <pub-id pub-id-type="doi">10.1121/1.1689343</pub-id></citation></ref>
<ref id="B23"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Freyman</surname> <given-names>R. L.</given-names></name> <name><surname>Helfer</surname> <given-names>K. S.</given-names></name> <name><surname>McCall</surname> <given-names>D. D.</given-names></name> <name><surname>Clifton</surname> <given-names>R. K.</given-names></name></person-group> (<year>1999</year>). <article-title>The role of perceived spatial separation in the unmasking of speech.</article-title> <source><italic>J. Acoust. Soc. Am.</italic></source> <volume>106</volume> <fpage>3578</fpage>&#x2013;<lpage>3588</lpage>. <pub-id pub-id-type="doi">10.1121/1.428211</pub-id></citation></ref>
<ref id="B24"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fridriksson</surname> <given-names>J.</given-names></name> <name><surname>Moss</surname> <given-names>J.</given-names></name> <name><surname>Davis</surname> <given-names>B.</given-names></name> <name><surname>Baylis</surname> <given-names>G. C.</given-names></name> <name><surname>Bonilha</surname> <given-names>L.</given-names></name> <name><surname>Rorden</surname> <given-names>C.</given-names></name></person-group> (<year>2008</year>). <article-title>Motor speech perception modulates the cortical language areas.</article-title> <source><italic>Neuroimage</italic></source> <volume>41</volume> <fpage>605</fpage>&#x2013;<lpage>613</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuroimage.2008.02.046</pub-id></citation></ref>
<ref id="B25"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Friston</surname> <given-names>K. J.</given-names></name> <name><surname>Bastos</surname> <given-names>A. M.</given-names></name> <name><surname>Pinotsis</surname> <given-names>D.</given-names></name> <name><surname>Litvak</surname> <given-names>V.</given-names></name></person-group> (<year>2015</year>). <article-title>LFP and oscillations&#x2014;what do they tell us?</article-title> <source><italic>Curr. Opin. Neurobiol.</italic></source> <volume>31</volume> <fpage>1</fpage>&#x2013;<lpage>6</lpage>. <pub-id pub-id-type="doi">10.1016/j.conb.2014.05.004</pub-id></citation></ref>
<ref id="B26"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gao</surname> <given-names>Y.-Y.</given-names></name> <name><surname>Cao</surname> <given-names>S.-Y.</given-names></name> <name><surname>Qu</surname> <given-names>T.-S.</given-names></name> <name><surname>Wu</surname> <given-names>X.-H.</given-names></name> <name><surname>Li</surname> <given-names>H.-F.</given-names></name> <name><surname>Zhang</surname> <given-names>J.-S.</given-names></name><etal/></person-group> (<year>2014</year>). <article-title>Voice-associated static face image releases speech from informational masking.</article-title> <source><italic>Psych. J.</italic></source> <volume>3</volume> <fpage>113</fpage>&#x2013;<lpage>120</lpage>. <pub-id pub-id-type="doi">10.1002/pchj.45</pub-id></citation></ref>
<ref id="B27"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Golumbic</surname> <given-names>E. M. Z.</given-names></name> <name><surname>Ding</surname> <given-names>N.</given-names></name> <name><surname>Bickel</surname> <given-names>S.</given-names></name> <name><surname>Lakatos</surname> <given-names>P.</given-names></name> <name><surname>Schevon</surname> <given-names>C. A.</given-names></name> <name><surname>McKhann</surname> <given-names>G. M.</given-names></name><etal/></person-group> (<year>2013</year>). <article-title>Mechanisms underlying selective neuronal tracking of attended speech at a &#x201C;Cocktail Party&#x201D;.</article-title> <source><italic>Neuron</italic></source> <volume>77</volume> <fpage>980</fpage>&#x2013;<lpage>991</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuron.2012.12.037</pub-id></citation></ref>
<ref id="B28"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Golumbic</surname> <given-names>E. M. Z.</given-names></name> <name><surname>Poeppel</surname> <given-names>D.</given-names></name> <name><surname>Schroeder</surname> <given-names>C. E.</given-names></name></person-group> (<year>2012</year>). <article-title>Temporal context in speech processing and attentional stream selection: a behavioral and neural perspective.</article-title> <source><italic>Brain Lang.</italic></source> <volume>122</volume> <fpage>151</fpage>&#x2013;<lpage>161</lpage>. <pub-id pub-id-type="doi">10.1016/j.bandl.2011.12.010</pub-id></citation></ref>
<ref id="B29"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hanslmayr</surname> <given-names>S.</given-names></name> <name><surname>Aslan</surname> <given-names>A.</given-names></name> <name><surname>Staudigl</surname> <given-names>T.</given-names></name> <name><surname>Klimesch</surname> <given-names>W.</given-names></name> <name><surname>Herrmann</surname> <given-names>C. S.</given-names></name> <name><surname>B&#x00E4;uml</surname> <given-names>K. H.</given-names></name></person-group> (<year>2007</year>). <article-title>Prestimulus oscillations predict visual perception performance between and within subjects.</article-title> <source><italic>Neuroimage</italic></source> <volume>37</volume> <fpage>1465</fpage>&#x2013;<lpage>1473</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuroimage.2007.07.011</pub-id></citation></ref>
<ref id="B30"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Helfer</surname> <given-names>K. S.</given-names></name></person-group> (<year>1997</year>). <article-title>Auditory and auditory-visual perception of clear and conversational speech.</article-title> <source><italic>J. Speech Lang. Hear. Res.</italic></source> <volume>40</volume> <fpage>432</fpage>&#x2013;<lpage>443</lpage>. <pub-id pub-id-type="doi">10.1044/jslhr.4002.432</pub-id></citation></ref>
<ref id="B31"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hickok</surname> <given-names>G.</given-names></name> <name><surname>Houde</surname> <given-names>J.</given-names></name> <name><surname>Rong</surname> <given-names>F.</given-names></name></person-group> (<year>2011</year>). <article-title>Sensorimotor integration in speech processing: computational basis and neural organization.</article-title> <source><italic>Neuron</italic></source> <volume>69</volume> <fpage>407</fpage>&#x2013;<lpage>422</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuron.2011.01.019</pub-id></citation></ref>
<ref id="B32"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hickok</surname> <given-names>G.</given-names></name> <name><surname>Poeppel</surname> <given-names>D.</given-names></name></person-group> (<year>2007</year>). <article-title>The cortical organization of speech processing.</article-title> <source><italic>Nat. Rev. Neurosci.</italic></source> <volume>8</volume> <fpage>393</fpage>&#x2013;<lpage>402</lpage>. <pub-id pub-id-type="doi">10.1038/nrn2113</pub-id></citation></ref>
<ref id="B33"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kerlin</surname> <given-names>J. R.</given-names></name> <name><surname>Shahin</surname> <given-names>A. J.</given-names></name> <name><surname>Miller</surname> <given-names>L. M.</given-names></name></person-group> (<year>2010</year>). <article-title>Attentional gain control of ongoing cortical speech representations in a &#x201C;cocktail party&#x201D;.</article-title> <source><italic>J. Neurosci.</italic></source> <volume>30</volume> <fpage>620</fpage>&#x2013;<lpage>628</lpage>. <pub-id pub-id-type="doi">10.1523/JNEUROSCI.3631-09.2010</pub-id></citation></ref>
<ref id="B34"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kong</surname> <given-names>Y. Y.</given-names></name> <name><surname>Mullangi</surname> <given-names>A.</given-names></name> <name><surname>Ding</surname> <given-names>N.</given-names></name></person-group> (<year>2014</year>). <article-title>Differential modulation of auditory responses to attended and unattended speech in different listening conditions.</article-title> <source><italic>Hear. Res.</italic></source> <volume>316</volume> <fpage>73</fpage>&#x2013;<lpage>81</lpage>. <pub-id pub-id-type="doi">10.1016/j.heares.2014.07.009</pub-id></citation></ref>
<ref id="B35"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lalor</surname> <given-names>E. C.</given-names></name> <name><surname>Foxe</surname> <given-names>J. J.</given-names></name></person-group> (<year>2010</year>). <article-title>Neural responses to uninterrupted natural speech can be extracted with precise temporal resolution.</article-title> <source><italic>Eur. J. Neurosci.</italic></source> <volume>31</volume> <fpage>189</fpage>&#x2013;<lpage>193</lpage>. <pub-id pub-id-type="doi">10.1111/j.1460-9568.2009.07055.x</pub-id></citation></ref>
<ref id="B36"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lewis</surname> <given-names>A. G.</given-names></name> <name><surname>Bastiaansen</surname> <given-names>M.</given-names></name></person-group> (<year>2015</year>). <article-title>A predictive coding framework for rapid neural dynamics during sentence-level language comprehension.</article-title> <source><italic>Cortex</italic></source> <volume>68</volume> <fpage>155</fpage>&#x2013;<lpage>168</lpage>. <pub-id pub-id-type="doi">10.1016/j.cortex.2015.02.014</pub-id></citation></ref>
<ref id="B37"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lewis</surname> <given-names>A. G.</given-names></name> <name><surname>Schoffelen</surname> <given-names>J. M.</given-names></name> <name><surname>Schriefers</surname> <given-names>H.</given-names></name> <name><surname>Bastiaansen</surname> <given-names>M.</given-names></name></person-group> (<year>2016</year>). <article-title>A predictive coding perspective on beta oscillations during sentence-level language comprehension.</article-title> <source><italic>Front. Hum. Neurosci.</italic></source> <volume>10</volume>:<issue>85</issue>. <pub-id pub-id-type="doi">10.3389/fnhum.2016.00085</pub-id></citation></ref>
<ref id="B38"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>L.</given-names></name> <name><surname>Daneman</surname> <given-names>M.</given-names></name> <name><surname>Qi</surname> <given-names>J. G.</given-names></name> <name><surname>Schneider</surname> <given-names>B. A.</given-names></name></person-group> (<year>2004</year>). <article-title>Does the information content of an irrelevant source differentially affect spoken word recognition in younger and older adults?</article-title> <source><italic>J. Exp. Psychol. Hum. Percept. Perform.</italic></source> <volume>30</volume> <fpage>1077</fpage>&#x2013;<lpage>1091</lpage>.</citation></ref>
<ref id="B39"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liberman</surname> <given-names>A. M.</given-names></name> <name><surname>Cooper</surname> <given-names>F. S.</given-names></name> <name><surname>Shankweiler</surname> <given-names>D. P.</given-names></name> <name><surname>Studdert-Kennedy</surname> <given-names>M.</given-names></name></person-group> (<year>1967</year>). <article-title>Perception of the speech code.</article-title> <source><italic>Psychol. Rev.</italic></source> <volume>74</volume> <fpage>431</fpage>&#x2013;<lpage>461</lpage>. <pub-id pub-id-type="doi">10.1037/h0020279</pub-id></citation></ref>
<ref id="B40"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liberman</surname> <given-names>A. M.</given-names></name> <name><surname>Delattre</surname> <given-names>P.</given-names></name> <name><surname>Cooper</surname> <given-names>F. S.</given-names></name></person-group> (<year>1952</year>). <article-title>The role of selected stimulus-variables in the perception of the unvoiced stop consonants.</article-title> <source><italic>Am. J. Psychol.</italic></source> <volume>65</volume> <fpage>497</fpage>&#x2013;<lpage>516</lpage>. <pub-id pub-id-type="doi">10.2307/1418032</pub-id></citation></ref>
<ref id="B41"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liberman</surname> <given-names>A. M.</given-names></name> <name><surname>Mattingly</surname> <given-names>I. G.</given-names></name></person-group> (<year>1985</year>). <article-title>The motor theory of speech perception revised.</article-title> <source><italic>Cognition</italic></source> <volume>21</volume> <fpage>1</fpage>&#x2013;<lpage>36</lpage>. <pub-id pub-id-type="doi">10.1016/0010-0277(85)90021-6</pub-id></citation></ref>
<ref id="B42"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Matthias</surname> <given-names>E.</given-names></name> <name><surname>Bublak</surname> <given-names>P.</given-names></name> <name><surname>M&#x00FC;ller</surname> <given-names>H. J.</given-names></name> <name><surname>Schneider</surname> <given-names>W. X.</given-names></name> <name><surname>Krummenacher</surname> <given-names>J.</given-names></name> <name><surname>Finke</surname> <given-names>K.</given-names></name></person-group> (<year>2010</year>). <article-title>The influence of alertness on spatial and nonspatial components of visual attention.</article-title> <source><italic>J. Exp. Psychol. Hum. Percept. Perform.</italic></source> <volume>36</volume> <fpage>38</fpage>&#x2013;<lpage>56</lpage>. <pub-id pub-id-type="doi">10.1037/a0017602</pub-id></citation></ref>
<ref id="B43"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Meister</surname> <given-names>I. G.</given-names></name> <name><surname>Wilson</surname> <given-names>S. M.</given-names></name> <name><surname>Deblieck</surname> <given-names>C.</given-names></name> <name><surname>Wu</surname> <given-names>A. D.</given-names></name> <name><surname>Iacoboni</surname> <given-names>M.</given-names></name></person-group> (<year>2007</year>). <article-title>The essential role of premotor cortex in speech perception.</article-title> <source><italic>Curr. Biol.</italic></source> <volume>17</volume> <fpage>1692</fpage>&#x2013;<lpage>1696</lpage>. <pub-id pub-id-type="doi">10.1016/j.cub.2007.08.064</pub-id></citation></ref>
<ref id="B44"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mesgarani</surname> <given-names>N.</given-names></name> <name><surname>Chang</surname> <given-names>E. F.</given-names></name></person-group> (<year>2012</year>). <article-title>Selective cortical representation of attended speaker in multi-talker speech perception.</article-title> <source><italic>Nature</italic></source> <volume>485</volume> <fpage>233</fpage>&#x2013;<lpage>236</lpage>. <pub-id pub-id-type="doi">10.1038/nature11020</pub-id></citation></ref>
<ref id="B45"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>O&#x2019;Sullivan</surname> <given-names>J. A.</given-names></name> <name><surname>Power</surname> <given-names>A. J.</given-names></name> <name><surname>Mesgarani</surname> <given-names>N.</given-names></name> <name><surname>Rajaram</surname> <given-names>S.</given-names></name> <name><surname>Foxe</surname> <given-names>J. J.</given-names></name> <name><surname>Shinn-Cunningham</surname> <given-names>B. G.</given-names></name><etal/></person-group> (<year>2014</year>). <article-title>Attentional selection in a cocktail party environment can be decoded from single-trial EEG.</article-title> <source><italic>Cereb. Cortex</italic></source> <volume>25</volume> <fpage>1697</fpage>&#x2013;<lpage>1706</lpage>. <pub-id pub-id-type="doi">10.1093/cercor/bht355</pub-id></citation></ref>
<ref id="B46"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pasley</surname> <given-names>B. N.</given-names></name> <name><surname>David</surname> <given-names>S. V.</given-names></name> <name><surname>Mesgarani</surname> <given-names>N.</given-names></name> <name><surname>Flinker</surname> <given-names>A.</given-names></name> <name><surname>Shamma</surname> <given-names>S. A.</given-names></name> <name><surname>Crone</surname> <given-names>N. E.</given-names></name><etal/></person-group> (<year>2012</year>). <article-title>Reconstructing speech from human auditory cortex.</article-title> <source><italic>PLoS Biol.</italic></source> <volume>10</volume>:<issue>e1001251</issue>. <pub-id pub-id-type="doi">10.1371/journal.pbio.1001251</pub-id></citation></ref>
<ref id="B47"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Piai</surname> <given-names>V.</given-names></name> <name><surname>Roelofs</surname> <given-names>A.</given-names></name> <name><surname>Rommers</surname> <given-names>J.</given-names></name> <name><surname>Dahlsl&#x00E4;tt</surname> <given-names>K.</given-names></name> <name><surname>Maris</surname> <given-names>E.</given-names></name></person-group> (<year>2015</year>). <article-title>Withholding planned speech is reflected in synchronized beta-band oscillations.</article-title> <source><italic>Front. Hum. Neurosci.</italic></source> <volume>9</volume>:<issue>549</issue>. <pub-id pub-id-type="doi">10.3389/fnhum.2015.00549</pub-id></citation></ref>
<ref id="B48"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Posner</surname> <given-names>M. I.</given-names></name> <name><surname>Petersen</surname> <given-names>S. E.</given-names></name></person-group> (<year>1990</year>). <article-title>The attention system of the human brain.</article-title> <source><italic>Annu. Rev. Neurosci.</italic></source> <volume>13</volume> <fpage>25</fpage>&#x2013;<lpage>42</lpage>. <pub-id pub-id-type="doi">10.1146/annurev.neuro.13.1.25</pub-id></citation></ref>
<ref id="B49"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Posner</surname> <given-names>M. I.</given-names></name> <name><surname>Petersen</surname> <given-names>S. E.</given-names></name></person-group> (<year>2012</year>). <article-title>The attention system of the human brain: 20 years after.</article-title> <source><italic>Annu. Rev. Neurosci.</italic></source> <volume>35</volume> <fpage>73</fpage>&#x2013;<lpage>89</lpage>. <pub-id pub-id-type="doi">10.1146/annurev-neuro-062111-150525</pub-id></citation></ref>
<ref id="B50"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Power</surname> <given-names>A. J.</given-names></name> <name><surname>Foxe</surname> <given-names>J. J.</given-names></name> <name><surname>Forde</surname> <given-names>E. J.</given-names></name> <name><surname>Reilly</surname> <given-names>R. B.</given-names></name> <name><surname>Lalor</surname> <given-names>E. C.</given-names></name></person-group> (<year>2012</year>). <article-title>At what time is the cocktail party? A late locus of selective attention to natural speech.</article-title> <source><italic>Eur. J. Neurosci.</italic></source> <volume>35</volume> <fpage>1497</fpage>&#x2013;<lpage>1503</lpage>. <pub-id pub-id-type="doi">10.1111/j.1460-9568.2012.08060.x</pub-id></citation></ref>
<ref id="B51"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Power</surname> <given-names>A. J.</given-names></name> <name><surname>Lalor</surname> <given-names>E. C.</given-names></name> <name><surname>Reilly</surname> <given-names>R. B.</given-names></name></person-group> (<year>2010</year>). <article-title>Endogenous auditory spatial attention modulates obligatory sensory activity in auditory cortex.</article-title> <source><italic>Cereb. Cortex</italic></source> <volume>21</volume> <fpage>1223</fpage>&#x2013;<lpage>1230</lpage>. <pub-id pub-id-type="doi">10.1093/cercor/bhq233</pub-id></citation></ref>
<ref id="B52"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pulverm&#x00FC;ller</surname> <given-names>F.</given-names></name> <name><surname>Huss</surname> <given-names>M.</given-names></name> <name><surname>Kherif</surname> <given-names>F.</given-names></name> <name><surname>del Prado Martin</surname> <given-names>F. M.</given-names></name> <name><surname>Hauk</surname> <given-names>O.</given-names></name> <name><surname>Shtyrov</surname> <given-names>Y.</given-names></name></person-group> (<year>2006</year>). <article-title>Motor cortex maps articulatory features of speech sounds.</article-title> <source><italic>Proc. Natl. Acad. Sci. U.S.A.</italic></source> <volume>103</volume> <fpage>7865</fpage>&#x2013;<lpage>7870</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.0509989103</pub-id></citation></ref>
<ref id="B53"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Roman</surname> <given-names>N.</given-names></name> <name><surname>Wang</surname> <given-names>D.</given-names></name> <name><surname>Brown</surname> <given-names>G. J.</given-names></name></person-group> (<year>2003</year>). <article-title>Speech segregation based on sound localization.</article-title> <source><italic>J. Acoust. Soc. Am.</italic></source> <volume>114</volume> <fpage>2236</fpage>&#x2013;<lpage>2252</lpage>. <pub-id pub-id-type="doi">10.1121/1.1610463</pub-id></citation></ref>
<ref id="B54"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Saarinen</surname> <given-names>T.</given-names></name> <name><surname>Jalava</surname> <given-names>A.</given-names></name> <name><surname>Kujala</surname> <given-names>J.</given-names></name> <name><surname>Stevenson</surname> <given-names>C.</given-names></name> <name><surname>Salmelin</surname> <given-names>R.</given-names></name></person-group> (<year>2015</year>). <article-title>Task-sensitive reconfiguration of corticocortical 6&#x2013;20 Hz oscillatory coherence in naturalistic human performance.</article-title> <source><italic>Hum. Brain Mapp.</italic></source> <volume>36</volume> <fpage>2455</fpage>&#x2013;<lpage>2469</lpage>. <pub-id pub-id-type="doi">10.1002/hbm.22784</pub-id></citation></ref>
<ref id="B55"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tadel</surname> <given-names>F.</given-names></name> <name><surname>Baillet</surname> <given-names>S.</given-names></name> <name><surname>Mosher</surname> <given-names>J. C.</given-names></name> <name><surname>Pantazis</surname> <given-names>D.</given-names></name> <name><surname>Leahy</surname> <given-names>R. M.</given-names></name></person-group> (<year>2011</year>). <article-title>Brainstorm: a user-friendly application for MEG/EEG analysis.</article-title> <source><italic>Comput. Intell. Neurosci.</italic></source> <volume>2011</volume>:<issue>879716</issue>. <pub-id pub-id-type="doi">10.1155/2011/879716</pub-id></citation></ref>
<ref id="B56"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Thorpe</surname> <given-names>S.</given-names></name> <name><surname>D&#x2019;Zmura</surname> <given-names>M.</given-names></name> <name><surname>Srinivasan</surname> <given-names>R.</given-names></name></person-group> (<year>2012</year>). <article-title>Lateralization of frequency-specific networks for covert spatial attention to auditory stimuli.</article-title> <source><italic>Brain Topogr.</italic></source> <volume>25</volume> <fpage>39</fpage>&#x2013;<lpage>54</lpage>. <pub-id pub-id-type="doi">10.1007/s10548-011-0186-x</pub-id></citation></ref>
<ref id="B57"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Todorovic</surname> <given-names>A.</given-names></name> <name><surname>Schoffelen</surname> <given-names>J. M.</given-names></name> <name><surname>van Ede</surname> <given-names>F.</given-names></name> <name><surname>Maris</surname> <given-names>E.</given-names></name> <name><surname>de Lange</surname> <given-names>F. P.</given-names></name></person-group> (<year>2015</year>). <article-title>Temporal expectation and attention jointly modulate auditory oscillatory activity in the beta band.</article-title> <source><italic>PLoS ONE</italic></source> <volume>10</volume>:<issue>e0120288</issue>. <pub-id pub-id-type="doi">10.1371/journal.pone.0120288</pub-id></citation></ref>
<ref id="B58"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>X. J.</given-names></name></person-group> (<year>2010</year>). <article-title>Neurophysiological and computational principles of cortical rhythms in cognition.</article-title> <source><italic>Physiol. Rev.</italic></source> <volume>90</volume> <fpage>1195</fpage>&#x2013;<lpage>1268</lpage>. <pub-id pub-id-type="doi">10.1152/physrev.00035.2008</pub-id></citation></ref>
<ref id="B59"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Weiss</surname> <given-names>S.</given-names></name> <name><surname>Mueller</surname> <given-names>H. M.</given-names></name></person-group> (<year>2012</year>). <article-title>&#x201C;Too many betas do not spoil the broth&#x201D;: the role of beta brain oscillations in language processing.</article-title> <source><italic>Front. Psychol.</italic></source> <volume>3</volume>:<issue>201</issue>. <pub-id pub-id-type="doi">10.3389/fpsyg.2012.00201</pub-id></citation></ref>
<ref id="B60"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wilson</surname> <given-names>S. M.</given-names></name> <name><surname>Iacoboni</surname> <given-names>M.</given-names></name></person-group> (<year>2006</year>). <article-title>Neural responses to non-native phonemes varying in producibility: evidence for the sensorimotor nature of speech perception.</article-title> <source><italic>Neuroimage</italic></source> <volume>33</volume> <fpage>316</fpage>&#x2013;<lpage>325</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuroimage.2006.05.032</pub-id></citation></ref>
<ref id="B61"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wilson</surname> <given-names>S. M.</given-names></name> <name><surname>Saygin</surname> <given-names>A. P.</given-names></name> <name><surname>Sereno</surname> <given-names>M. I.</given-names></name> <name><surname>Iacoboni</surname> <given-names>M.</given-names></name></person-group> (<year>2004</year>). <article-title>Listening to speech activates motor areas involved in speech production.</article-title> <source><italic>Nat. Neurosci.</italic></source> <volume>7</volume> <fpage>701</fpage>&#x2013;<lpage>702</lpage>. <pub-id pub-id-type="doi">10.1038/nn1263</pub-id></citation></ref>
<ref id="B62"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Womelsdorf</surname> <given-names>T.</given-names></name> <name><surname>Fries</surname> <given-names>P.</given-names></name></person-group> (<year>2007</year>). <article-title>The role of neuronal synchronization in selective attention.</article-title> <source><italic>Curr. Opin. Neurobiol.</italic></source> <volume>17</volume> <fpage>154</fpage>&#x2013;<lpage>160</lpage>. <pub-id pub-id-type="doi">10.1016/j.conb.2007.02.002</pub-id></citation></ref>
<ref id="B63"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wu</surname> <given-names>Z.-M.</given-names></name> <name><surname>Chen</surname> <given-names>M.-L.</given-names></name> <name><surname>Wu</surname> <given-names>X.-H.</given-names></name> <name><surname>Li</surname> <given-names>L.</given-names></name></person-group> (<year>2014</year>). <article-title>Interaction between auditory system and motor system in speech perception.</article-title> <source><italic>Neurosci. Bull.</italic></source> <volume>30</volume> <fpage>490</fpage>&#x2013;<lpage>496</lpage>. <pub-id pub-id-type="doi">10.1007/s12264-013-1428-6</pub-id></citation></ref>
<ref id="B64"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yang</surname> <given-names>Z.</given-names></name> <name><surname>Chen</surname> <given-names>J.</given-names></name> <name><surname>Huang</surname> <given-names>Q.</given-names></name> <name><surname>Wu</surname> <given-names>X.</given-names></name> <name><surname>Wu</surname> <given-names>Y.</given-names></name> <name><surname>Schneider</surname> <given-names>B. A.</given-names></name><etal/></person-group> (<year>2007</year>). <article-title>The effect of voice cuing on releasing Chinese speech from informational masking.</article-title> <source><italic>Speech Commun.</italic></source> <volume>49</volume> <fpage>892</fpage>&#x2013;<lpage>904</lpage>. <pub-id pub-id-type="doi">10.1097/AUD.0b013e3181db6dc2</pub-id></citation></ref>
</ref-list>
<fn-group>
<fn id="fn01"><label>1</label><p><ext-link ext-link-type="uri" xlink:href="http://neuroimage.usc.edu/brainstorm/">http://neuroimage.usc.edu/brainstorm/</ext-link></p></fn>
</fn-group>
</back>
</article>