<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="brief-report">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Neurosci.</journal-id>
<journal-title>Frontiers in Neuroscience</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Neurosci.</abbrev-journal-title>
<issn pub-type="epub">1662-453X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fnins.2025.1534425</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Neuroscience</subject>
<subj-group>
<subject>Perspective</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Understanding speech in &#x0201C;noise&#x0201D; or free energy minimization in the soundscapes of the anthropocene</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes" equal-contrib="yes">
<name><surname>Strauss</surname> <given-names>Daniel J.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x0002A;</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/240319/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Francis</surname> <given-names>Alexander L.</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/211647/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Sch&#x000E4;fer</surname> <given-names>Zeinab</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1072618/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Latzel</surname> <given-names>Matthias</given-names></name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Corona&#x02013;Strauss</surname> <given-names>Farah I.</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Launer</surname> <given-names>Stefan</given-names></name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/1336777/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Systems Neuroscience and Neurotechnology Unit, Faculty of Medicine, Saarland University, and School of Engineering, htw saar</institution>, <addr-line>Homburg/Saar</addr-line>, <country>Germany</country></aff>
<aff id="aff2"><sup>2</sup><institution>Speech Perception and Cognitive Effort Lab, Department of Speech, Language and Hearing Sciences, Purdue University</institution>, <addr-line>West Lafayette, IN</addr-line>, <country>United States</country></aff>
<aff id="aff3"><sup>3</sup><institution>Department of Computer Science, Aarhus University</institution>, <addr-line>Aarhus</addr-line>, <country>Denmark</country></aff>
<aff id="aff4"><sup>4</sup><institution>Sonova AG</institution>, <addr-line>St&#x000E4;fa</addr-line>, <country>Switzerland</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: Achim Klug, University of Colorado Anschutz Medical Campus, United States</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Darius Parvizi-Wayne, University College London, United Kingdom</p></fn>
<corresp id="c001">&#x0002A;Correspondence: Daniel J. Strauss <email>daniel.strauss&#x00040;uni-saarland.de</email></corresp>
<fn fn-type="equal" id="fn001"><p>&#x02020;These authors have contributed equally to this work</p></fn></author-notes>
<pub-date pub-type="epub">
<day>14</day>
<month>03</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>19</volume>
<elocation-id>1534425</elocation-id>
<history>
<date date-type="received">
<day>25</day>
<month>11</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>19</day>
<month>02</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x000A9; 2025 Strauss, Francis, Sch&#x000E4;fer, Latzel, Corona&#x02013;Strauss and Launer.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Strauss, Francis, Sch&#x000E4;fer, Latzel, Corona&#x02013;Strauss and Launer</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<p>Listening to speech in the presence of irrelevant sounds is ubiquitous in the modern world, but is generally acknowledged to be both effortful and unpleasant. Here we argue that this problem arises largely in circumstances that our human auditory system has not evolved to accommodate. The soundscapes of the Anthropocene are frequently characterized by an overabundance of sound sources, the vast majority of which are functionally irrelevant to a given listener. The problem of listening to speech in such environments must be solved by an auditory system that is not optimized for this task. Building on our previous work linking attention to effortful listening and incorporating an active inference approach, we argue that the answers to these questions have implications not just for the study of human audition. They are also significant for the development and broad awareness of hearing aids and cochlear implants, as well as other auditory technologies such as earbuds, immersive auditory environments, and systems for human-machine interaction.</p></abstract>
<kwd-group>
<kwd>hearing</kwd>
<kwd>evolution</kwd>
<kwd>noise</kwd>
<kwd>free energy principle</kwd>
<kwd>attention</kwd>
<kwd>listening effort</kwd>
<kwd>speech</kwd>
</kwd-group>
<counts>
<fig-count count="2"/>
<table-count count="0"/>
<equation-count count="1"/>
<ref-count count="71"/>
<page-count count="8"/>
<word-count count="6431"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Auditory Cognitive Neuroscience</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<title>1 Introduction</title>
<p>The Anthropocene, a term referring to our current era within the geological epoch of the Holocene, is marked by human activities that radically alter our soundscapes, see Habib et al. (<xref ref-type="bibr" rid="B29">2007</xref>); Swaddle et al. (<xref ref-type="bibr" rid="B68">2015</xref>); Slabbekoorn (<xref ref-type="bibr" rid="B59">2018</xref>). Modern soundscapes incorporate sounds produced by large numbers of people, traffic, machinery, phones, radios, televisions, etc. but also the indirect sounds produced by reverberations from installed man-made objects and infrastructure. In the acoustically rich and complex soundscape of the Anthropocene, the amount and variety of unimportant and unwanted sound subsumed as &#x0201D;noise&#x0201D; has steadily increased (Habib et al., <xref ref-type="bibr" rid="B29">2007</xref>; Slabbekoorn, <xref ref-type="bibr" rid="B59">2018</xref>).</p>
<p>The development of transportation systems, encompassing vehicles, aircraft, and maritime vessels, has extended noise pollution into previously remote areas. Industrialization, through activities such as mining and energy production, further exacerbates noise pollution. These drastic changes in the soundscape clearly impact wildlife but also human habitats (Habib et al., <xref ref-type="bibr" rid="B29">2007</xref>; Swaddle et al., <xref ref-type="bibr" rid="B68">2015</xref>; Slabbekoorn, <xref ref-type="bibr" rid="B59">2018</xref>). Similarly, anthropogenic climate change has indirect effects on natural soundscapes by altering weather patterns, habitats, and animal behavior, but also reshapes regional sound compositions critical to human wellbeing (Lorenzi et al., <xref ref-type="bibr" rid="B36">2023</xref>).</p>
<p>Overall, the Anthropocene has ushered in a notable escalation in human-induced noise pollution, necessitating concerted efforts in sound monitoring, regulatory measures, and the adoption of quieter technologies to mitigate its adverse effects on ecosystems, wildlife, and human health.</p>
<p>With respect to anthropogenic changes that directly affect human listening, the most significant of these is urbanization, which has introduced heightened levels of human-generated sounds, and concentrated humans together in unprecedented numbers.</p>
<p>Thus, there has been a change in the soundscape, but also in what we need to get from the soundscape. In earlier eras such as the Pleistocene epoch, the time period in which modern humans evolved (Tooby and Cosmides, <xref ref-type="bibr" rid="B70">1992</xref>), listening to the non-human world was much more important for survival, for example in order to avoid threats and achieve goals (hunting, gathering, etc.). Sudden and/or high-intensity, &#x0201D;attention-grabbing&#x0201D; sounds were likely to be important for survival, potentially signaling significant changes in the immediate environment, and thus early auditory systems (i.e., those inherited by early hominins) would likely have already evolved to treat them with priority. Surrounding speech was likely produced by known individuals, and was likely to be important for social interaction and, ultimately, survival. In the Anthropocene, we use our senses in a very different way than in the deep past. We are (mostly) not threatened by anything that makes sound but is not human. A sudden sound like the breaking of a nearby branch, or the warning call of a bird or small mammal, is functionally irrelevant to most humans in the Anthropocene. Even the sound of a neighbor&#x00027;s car door slamming, or the clatter of glasses in a restaurant kitchen, while attention-demanding, is generally functionally irrelevant. The relative significance of exogenously-directed auditory attention has changed radically in many modern contexts, typically with far less relevance to immediate survival except, quite notably, in the case of avoiding traffic. And yet, despite the decline in the fitness benefit of orienting toward sudden, loud, warning-like sounds, there are ever more sounds in the environment that may cause an involuntary switch of attention&#x02014;for instance, the squeal of a tram, the slam of a door, a car horn. Even though none of these sounds may be relevant, they all exhibit acoustic properties that make them attentionally demanding, i.e., distracting, and the repeated capture of attention by irrelevant sounds quickly becomes annoying and stressful.</p>
<p>Using careful but extremely incomplete phylogenetic information from Ackermann et al. (<xref ref-type="bibr" rid="B2">2014</xref>); Gintis (<xref ref-type="bibr" rid="B25">2011</xref>); Chen and Wiens (<xref ref-type="bibr" rid="B11">2020</xref>), we have loosely sketched the changing complexity of soundscapes in urbanized areas (see also Slabbekoorn, <xref ref-type="bibr" rid="B59">2018</xref>) and auditory capacity in a conceptual time relation. The point is to show that evolutionary timescales and the timescale of changing soundscapes of the Anthropocene differ vividly. Here, we are deliberately vague as to what exactly constitutes auditory capacity, though for the present, we can roughly define it as the ability for sound processing and production in acoustic communication in vertebrates (in the Mesozoic Era) and our lineage in primates later in the Cenozoic Era. In primates, sophisticated vocal learning, speech, and comprehensive musicality seem to be specific to humans (<italic>Homo sapiens</italic>) (Dichter et al., <xref ref-type="bibr" rid="B17">2018</xref>; Aboitiz, <xref ref-type="bibr" rid="B1">2018</xref>; Patel, <xref ref-type="bibr" rid="B47">2021</xref>), though some aspects of spoken language were likely present in other now-extinct <italic>Homo</italic> species, including <italic>neanderthalis</italic> (Conde-Valverde et al., <xref ref-type="bibr" rid="B14">2021</xref>) and <italic>erectus</italic> (Swedell and Plummer, <xref ref-type="bibr" rid="B69">2019</xref>; Everett, <xref ref-type="bibr" rid="B20">2017</xref>). In fact, auditory capacity has not changed much since Homo sapiens appeared (approx. 300,000&#x02013;200,000 years ago), see Everett (<xref ref-type="bibr" rid="B20">2017</xref>). For a more general review on the development of acoustic communication and auditory capacity within mammals, we refer to Grothe et al. (<xref ref-type="bibr" rid="B28">2004</xref>); Sterbing-D&#x00027;Angelo (<xref ref-type="bibr" rid="B61">2009</xref>); Ackermann et al. (<xref ref-type="bibr" rid="B2">2014</xref>); Manley (<xref ref-type="bibr" rid="B37">2017</xref>); Chen and Wiens (<xref ref-type="bibr" rid="B11">2020</xref>).</p>
<p>Since the arrival of Homo sapiens, beside genetic drifts (Star and Spencer, <xref ref-type="bibr" rid="B60">2013</xref>) and epigenetic factors (Ashe and Colot, <xref ref-type="bibr" rid="B3">2021</xref>), cultural evolution has far outpaced any changes due to evolution. Cultural change has altered our environmental soundscapes in urbanized areas drastically, in particular since the industrial revolution and the exponentially increasing use of technology, see <xref ref-type="fig" rid="F1">Figure 1</xref>. Note that this discussion does not include the acceleration of human adaptive evolution due to gene-environment interactions and a gene-culture co-evolution (e.g., see Hawks et al., <xref ref-type="bibr" rid="B31">2007</xref>; Gintis, <xref ref-type="bibr" rid="B25">2011</xref>) or factors which are not directly related to the &#x0201C;auditory capacity&#x0201D;. It is also worth emphasizing that, in the following section, we focus on deeply anchored auditory attention mechanisms in Homo sapiens and not on culturally dependent learning and adaptation mechanisms of the attention system within our species (Boduroglu and Shah, <xref ref-type="bibr" rid="B4">2017</xref>; Jurkat et al., <xref ref-type="bibr" rid="B34">2020</xref>) in the modern age.</p>
<fig id="F1" position="float">
<label>Figure 1</label>
<caption><p>Conceptual sketch of the changing complexity of soundscapes and the auditory capacity in the Mesozoic and Cenozoic Era. Whereas the auditory capacity is rather constant since the Homo sapiens appeared, the soundscapes of urbanized areas changed drastically since the industrial and technological revolution.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnins-19-1534425-g0001.tif"/>
</fig>
</sec>
<sec id="s2">
<title>2 Conceptual model for understanding speech in noise</title>
<p>So how do attention mechanisms in hearing from the times of hunters and gatherers fit the modern world? Listening occurs in a soundscape that is always present. Outside the hearing clinic or laboratory it is virtually impossible to find a single acoustic event occurring in isolation. The ability to subset sensory space to better distinguish relevant from irrelevant phenomena is ecologically essential (Stevens, <xref ref-type="bibr" rid="B62">2013</xref>; Lev-Ari et al., <xref ref-type="bibr" rid="B35">2022</xref>; Bruner and Colom, <xref ref-type="bibr" rid="B8">2022</xref>), and plays a particularly important role in social primates (Sch&#x000FC;lke et al., <xref ref-type="bibr" rid="B56">2020</xref>). However, unlike in vision, humans and apes do not significantly move their pinnae toward sounds, even though we retained a vestigial pinna-orienting system that has persisted as a &#x0201C;neural fossil&#x0201D; within in the brain for about 25 million years (Hackley, <xref ref-type="bibr" rid="B30">2015</xref>; Strauss et al., <xref ref-type="bibr" rid="B64">2020</xref>). So, we cannot move our ears as we can move our eyes to focus on a part of the visual scene. We do not have an auditory fovea, at least not in a physical sense. Thus in listening, unlike in vision, the neural representation of the physical world remains unchanged by the peripheral sensor itself. In order to tune in to a particular auditory object within a soundscape (auditory scene) we need to use attention (Shinn-Cunningham, <xref ref-type="bibr" rid="B58">2008</xref>).</p>
<p>Here, to distinguish between the target sound and the unwanted sound, i.e., the noise, we can employ the binary figure (target) and background (noise) principle from vision, see Marr (<xref ref-type="bibr" rid="B38">1982</xref>). The idea of figure and background is very important because there is always a whole soundscape. What distinguishes figure from ground are the &#x0201C;goals&#x0201D; of the listener. Note that, although typically such &#x0201C;goals&#x0201D; are considered in terms of explicit representations (e.g. &#x0201C;I want to listen to that person, not this one&#x0201D;), we prefer to consider them more broadly, even including such &#x0201C;corporeal beliefs&#x0201D; as &#x0201C;I want to avoid danger&#x0201D; (see Parvizi-Wayne, <xref ref-type="bibr" rid="B46">2024</xref>). Thus, a voice might shift from background to foreground either because I choose to attend to it, or because it has become louder and higher pitched, as it might if the speaker is angry and potentially becoming a threat.</p>
<p>In the pre-Anthropocene era, many sounds in the soundscape were crucial for survival, and natural selection would over time have tuned our senses, and our attention, to better respond to them. In the sense of Parvizi-Wayne (<xref ref-type="bibr" rid="B46">2024</xref>), rapid response to sudden, loud sounds, for example, should be deeply entrenched in the predictive model that guides what we pay attention to. Thus, an animal hunting or foraging in a small group might be listening primarily to the environment, highly sensitive to any change that might signal the approach of a threat or loss of an opportunity. In species that forage in groups, including both baboons and chimpanzees, this includes listening not just to the sounds of the environment, but also to one&#x00027;s friends and neighbors, who are often both allies and potential rivals. That is, although we often focus on the survival benefits of listening in the context of predator/prey interactions, it seems likely that, in a social animal the ability to listen to relevant communication (directed both to the listener and to others) is at least equally significant (Sch&#x000FC;lke et al., <xref ref-type="bibr" rid="B56">2020</xref>).</p>
<sec>
<title>2.1 Segregating the target from &#x0201C;noise&#x0201D;</title>
<p>There is a strong link between different modes of attention and effortful listening to speech as a target in noise (see Strauss and Francis, <xref ref-type="bibr" rid="B66">2017</xref>). Typically, one starts with the classic taxonomy of exogenous attention (bottom-up, automatic, unconscious) and endogenous attention (voluntary/top-down, goal-directed) in sensory processing (see M&#x000FC;ller and Rabbitt, <xref ref-type="bibr" rid="B41">1989</xref>; Jigo et al., <xref ref-type="bibr" rid="B33">2021</xref>; Ren et al., <xref ref-type="bibr" rid="B49">2021</xref>) frequently applied in auditory scene analysis (Bregman, <xref ref-type="bibr" rid="B7">1990</xref>). However, understanding the distribution of attention does not necessarily depend on a strict division of these concepts. For example, following Parvizi-Wayne (<xref ref-type="bibr" rid="B46">2024</xref>), the exogenous attraction of attention by an external stimulus (e.g., a sudden, loud sound) can still be understood as the deployment of attention toward a stimulus on the basis of goals, albeit goals that may be deeply entrenched in the predictive, hierarchically organized model of the environment (e.g., &#x0201C;Identify the source of sudden, loud sounds&#x0201D;) as a result of millennia of natural selection. Thus, for example, attentional focus can be modeled by a continuous (probabilistic) stream selection model depending on weights related to exogenous and endogenous processes (Trenado et al., <xref ref-type="bibr" rid="B71">2009</xref>; Strauss et al., <xref ref-type="bibr" rid="B65">2010</xref>) or by a single-agent model when using active inference, which does not need a dichotomy of these attention concepts (see Parvizi-Wayne, <xref ref-type="bibr" rid="B46">2024</xref> and below). In either case, the probabilistic selection scheme in Trenado et al. (<xref ref-type="bibr" rid="B71">2009</xref>) is akin to the biased competition model in visual perception (Desimone and Duncan, <xref ref-type="bibr" rid="B16">1995</xref>); see also Shinn-Cunningham (<xref ref-type="bibr" rid="B58">2008</xref>) for an adaptation of biased competition to the auditory modality and Strauss et al. (<xref ref-type="bibr" rid="B65">2010</xref>) for a mapping to effortful listening. The two-competitor as well as taxonomic models of attention in effortful listening are summarized in <xref ref-type="fig" rid="F2">Figure 2</xref>, showing a typical cocktail party situation where two people are having a conversation (target) in a noisy background (see also below).</p>
<fig id="F2" position="float">
<label>Figure 2</label>
<caption><p>The relation between attention and effort in listening: <bold>(A)</bold> Model of the attentional competition between auditory target and background based on a modified version of the two-competitor model in Strauss et al. (<xref ref-type="bibr" rid="B67">2024b</xref>). It is a dynamical model in which effort varies in the 2D plane, and we have to invest effort to keep the attentional focus vector on the target. <bold>(B)</bold> The associated effort of the target processing when employing the taxonomic model of attention in effortful listening in Strauss and Francis (<xref ref-type="bibr" rid="B66">2017</xref>). Here, the exerted attentional effort (Sarter et al., <xref ref-type="bibr" rid="B51">2006</xref>; Strauss and Francis, <xref ref-type="bibr" rid="B66">2017</xref>) is driven by the internal demand d<sub><italic>i</italic></sub> (e.g., requiring attention to working memory objects to make sense of complicated spoken sentences) and the external demand <italic>d</italic><sub><italic>e</italic></sub> (e.g., requiring attention to perform stream segregation to separate the target from the &#x0201C;noise&#x0201D;), see Strauss and Francis (<xref ref-type="bibr" rid="B66">2017</xref>); Strauss et al. (<xref ref-type="bibr" rid="B67">2024b</xref>) for details.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnins-19-1534425-g0002.tif"/>
</fig>
<p>The predictive (generative) hierarchically organized model of the environment employed in Feldman and Friston (<xref ref-type="bibr" rid="B21">2010</xref>); Clark (<xref ref-type="bibr" rid="B12">2013</xref>); Parvizi-Wayne (<xref ref-type="bibr" rid="B46">2024</xref>) might, for the sake of simplicity in structure, be mathematically seen as a generative spatiotemporal scale space, i.e., a nested structure encompassing multiple spatiotemporal scales of the (predicted) environment. This scale space covers predictions about mostly fast and nearby events at an evolutionarily deeply entrenched core of the model. The further we move away from this core, the more complex and forward-thinking (in time and space) these predictions might be. Building on Feldman and Friston (<xref ref-type="bibr" rid="B21">2010</xref>) that treats biased competition in terms of an active-inference framework, Parvizi-Wayne (<xref ref-type="bibr" rid="B46">2024</xref>) draws on Friston&#x00027;s free energy minimization theory (Friston, <xref ref-type="bibr" rid="B22">2010</xref>) to represent the often distinct ideas of exogenous and endogenous attention in a unitary, active inference framework. As this framework comprehensively supports a mapping of the relation between perception and action in effortful listening, let us take a more formal look at its structure.</p>
</sec>
<sec>
<title>2.2 Free energy principle in effortful listening</title>
<p>In real-world scenarios such as the cocktail party in <xref ref-type="fig" rid="F2">Figure 2</xref>, listening is not happening in isolation, e.g., deprived of the other senses. Regularities across senses, e.g., integrating lip reading or posture with listening, are crucial for understanding speech in noise, see, e.g., Rosenblum (<xref ref-type="bibr" rid="B50">2008</xref>) and also our discussion of relation between perception and action below. However, to avoid an over-generalization, we focus on the auditory modality in the following formal discussion, assuming implicitly that the brain has an internal scale space representation of the entire multisensory environment. We formulate the free energy according to Parr et al. (<xref ref-type="bibr" rid="B45">2022</xref>) as follows for our auditory setting; The internal state of the generative scale space model corresponds to a distribution <italic>Q</italic>(<italic>s</italic>), which captures the brain&#x00027;s expectations and prior beliefs about the acoustic scene <italic>s</italic> (including both figure and background across spatiotemporal scales). Let <italic>o</italic> represent sensory input, such as the mixture of sounds in a cocktail party scenario. The free energy <italic>F</italic>(<italic>Q, o</italic>) can be expressed as:</p>
<disp-formula id="E1"><mml:math id="M1"><mml:mtable columnalign="left"><mml:mtr><mml:mtd><mml:mstyle displaystyle="true"><mml:munder class="msub"><mml:mrow><mml:mstyle displaystyle="true"><mml:munder accentunder="false"><mml:mrow><mml:mi>F</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>Q</mml:mi><mml:mo>,</mml:mo><mml:mi>o</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo>&#x0FE38;</mml:mo></mml:munder></mml:mstyle></mml:mrow><mml:mrow><mml:mtext class="textrm" mathvariant="normal">free energy</mml:mtext></mml:mrow></mml:munder></mml:mstyle><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:munder class="msub"><mml:mrow><mml:mstyle displaystyle="true"><mml:munder accentunder="false"><mml:mrow><mml:msub><mml:mrow><mml:mi>E</mml:mi></mml:mrow><mml:mrow><mml:mi>Q</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>s</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow></mml:msub><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mo>-</mml:mo><mml:mo class="qopname">ln</mml:mo><mml:mtext>&#x000A0;</mml:mtext><mml:mi>p</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>s</mml:mi><mml:mo>,</mml:mo><mml:mi>o</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo>]</mml:mo></mml:mrow></mml:mrow><mml:mo>&#x0FE38;</mml:mo></mml:munder></mml:mstyle></mml:mrow><mml:mrow><mml:mtext class="textrm" mathvariant="normal">expectation term</mml:mtext></mml:mrow></mml:munder></mml:mstyle><mml:mo>-</mml:mo><mml:mstyle displaystyle="true"><mml:munder class="msub"><mml:mrow><mml:mstyle displaystyle="true"><mml:munder accentunder="false"><mml:mrow><mml:mi>H</mml:mi><mml:mrow><mml:mo>[</mml:mo><mml:mrow><mml:mi>Q</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>s</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo>]</mml:mo></mml:mrow></mml:mrow><mml:mo>&#x0FE38;</mml:mo></mml:munder></mml:mstyle></mml:mrow><mml:mrow><mml:mtext class="textrm" mathvariant="normal">entropy term</mml:mtext></mml:mrow></mml:munder></mml:mstyle><mml:mo>&#x02265;</mml:mo><mml:mstyle displaystyle="true"><mml:munder class="msub"><mml:mrow><mml:mstyle displaystyle="true"><mml:munder accentunder="false"><mml:mrow><mml:mo>-</mml:mo><mml:mo class="qopname">ln</mml:mo><mml:mtext>&#x000A0;</mml:mtext><mml:mi>p</mml:mi><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mi>o</mml:mi></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mo>&#x0FE38;</mml:mo></mml:munder></mml:mstyle></mml:mrow><mml:mrow><mml:mtext class="textrm" mathvariant="normal">surprise</mml:mtext></mml:mrow></mml:munder></mml:mstyle></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula>
<p>In this formulation, <italic>Q</italic>(<italic>s</italic>) serves as both the approximate posterior distribution and the internal state itself, representing the brain&#x00027;s model of the acoustic scene. The first term (the negative expected log joint probability, also known as the energy) measures how well the brain&#x00027;s internal model predicts auditory input given its current beliefs. The second term <italic>H</italic>[<italic>Q</italic>(<italic>s</italic>)] which measures the entropy of this distribution, quantifies uncertainty in the brain&#x00027;s beliefs about the acoustic environment. A higher entropy value implies greater uncertainty in the listener&#x00027;s beliefs, making it more receptive to new sensory observations. This adaptability aligns with the Free Energy Principle, which suggests that the brain minimizes free energy by reducing uncertainty and refining its internal model to optimize perception and action in a dynamic environment (Friston, <xref ref-type="bibr" rid="B22">2010</xref>; Parr et al., <xref ref-type="bibr" rid="B45">2022</xref>).</p>
<p>Minimizing free energy requires balancing consistency with the generative model (expectation term) while maintaining appropriate uncertainty through the entropy maximization. In the absence of precise prior beliefs, this formulation follows Jaynes&#x00027;s maximum entropy principle, suggesting the perceptual system should maintain maximum uncertainty about hidden states when information is limited (i.e., higher entropy enables more flexible adjustment of internal models). Importantly, the inequality in the free energy formulation indicates that free energy always provides an upper bound on &#x0201C;surprise&#x0201D;&#x02014;the unexpectedness of sensory input. By minimizing free energy, the perceptual system indirectly minimizes surprise, making auditory inputs more predictable and allowing for more effective processing of complex acoustic scenes.</p>
<p>This interaction of perception and action might occur quickly and automatically at the evolutionarily deeply entrenched core of the scale space model but also more slowly due to the engagement of more complex and thoughtful schemas as we move away from the core. As predictions and expectations drive attention (e.g., see Grossberg, <xref ref-type="bibr" rid="B27">2005</xref>; Strauss et al., <xref ref-type="bibr" rid="B65">2010</xref>; Clark, <xref ref-type="bibr" rid="B12">2013</xref>; Parvizi-Wayne, <xref ref-type="bibr" rid="B46">2024</xref>), fast and automatic links between perception and action are associated with the &#x0201C;exogenous weights&#x0201D; and slower, reasoning based loops with the &#x0201C;endogenous weights&#x0201D; in the classic terminology (Strauss and Francis, <xref ref-type="bibr" rid="B66">2017</xref>; Strauss et al., <xref ref-type="bibr" rid="B67">2024b</xref>). Parvizi-Wayne&#x00027;s model (Parvizi-Wayne, <xref ref-type="bibr" rid="B46">2024</xref>) emphasizes in this context the importance of &#x0201C;precision weight optimization&#x0201D; over &#x0201C;precision optimization&#x0201D;. This involves not only making predictions across spatiotemporal scales more accurate but also determining the importance (weight) of different pieces of information in their scale space representations (i.e., the multiscale prediction error of perceptual beliefs). If these pieces of information are representations of predicted auditory objects (Shinn-Cunningham, <xref ref-type="bibr" rid="B58">2008</xref>), the following discussion will clarify how this model maps the cocktail party situation illustrated in <xref ref-type="fig" rid="F2">Figure 2</xref>. As we will see, these weights depend on the auditory scene, the context, and the generative model including learned experiences. Precision weight optimization would also map enhanced representations of attended auditory objects along the hearing path due to attention &#x0201C;gain&#x0201D; (or &#x0201C;noise suppression&#x0201D;) neural mechanisms (Strauss et al., <xref ref-type="bibr" rid="B63">2024a</xref>), see Parvizi-Wayne (<xref ref-type="bibr" rid="B46">2024</xref>) for more detailed discussions.</p>
<p>Minimizing free energy and surprise directly answers the question of what motivates us to follow a &#x0201C;listening goal&#x0201D;. This &#x0201C;motivational aspect&#x0201D; involves long-term and reasoning-based minimization of (negative) surprises across spatiotemporal scales. For instance, assume the man in <xref ref-type="fig" rid="F2">Figure 2</xref> is telling the woman what changes are being planned in the dean&#x00027;s office for next week or at the federal level regarding energy prices next winter. As she does not want to encounter surprises in these matters, she is motivated to exert attentional effort in the conversation (Strauss and Francis, <xref ref-type="bibr" rid="B66">2017</xref>). However, free energy minimization also applies to the here and now, providing an almost instantaneous and automatic analysis of sudden loud sounds in the acoustic scene, causing a free energy spike (Parvizi-Wayne, <xref ref-type="bibr" rid="B46">2024</xref>). Consider an abrupt background sound like laughter or clinking glasses at the cocktail party in <xref ref-type="fig" rid="F2">Figure 2</xref>. These automatic processes stem from an evolutionarily deeply entrenched core of the scale space model, securing survival in the present moment. No matter how interesting the conversation about future events is, these acoustically salient events compete for our attention, distract us from the conversation, and cause increased attentional effort to follow the conversation and minimize free energy about future events (called exogenous override in Strauss et al. (<xref ref-type="bibr" rid="B67">2024b</xref>)). Here it does not matter that clinking glasses do not need our attention at the cocktail party. As we have stated before, we are not evolutionarily optimized for cocktail parties or other features of modern urban environments. In the modern world, there is an abundance of these free energy spikes caused by sounds addressing the &#x0201C;survival core&#x0201D; of our generative hierarchical model of the environment and also an abundance of acoustic information that might be worth following when minimizing surprises. It is also important to note that attentional shifts in noisy environments can arise from individual priors (i.e., learned experiences) embedded in the brain&#x00027;s internal model, causing one person to focus on a background stimulus linked to past experiences, while another perceives it as irrelevant. For instance, the ringtone of a cell phone might capture more attention if it is the same as one&#x00027;s own, or the sound of a falling tablet could be associated with a threatening situation at a previous cocktail party based on individual experience. Turning to our major theme, the Anthropocene is inducing more free energy in listening situations in multiple ways, vastly increasing the effort of maintaining simple conversations at a cocktail party, let alone in Times Square at rush hour.</p>
</sec>
</sec>
<sec id="s3">
<title>3 Discussion and technological implications</title>
<p>Technological advances have significantly altered even modern soundscapes within a single human lifetime (Habib et al., <xref ref-type="bibr" rid="B29">2007</xref>; Swaddle et al., <xref ref-type="bibr" rid="B68">2015</xref>; Slabbekoorn, <xref ref-type="bibr" rid="B59">2018</xref>). Cultural evolution, driven by technology, has far outpaced biological evolution (Boyd et al., <xref ref-type="bibr" rid="B6">2013</xref>), leaving us with a sensory processing and perceptual system naturally equipped for environments vastly different from our world today (e.g., see Gazzaley and Rosen, <xref ref-type="bibr" rid="B23">2016</xref>). Our arguments allow us to look at the idea of restoring hearing to its &#x0201C;natural&#x0201D; state from a new angle as our auditory system evolved for vastly different environments. Rather than simply restoring hearing, augmenting hearing through technologies such as noise suppression and directional microphones enables a technological adaptation to Anthropocene soundscapes rather than simply restoring Pleistocene capabilities (see <xref ref-type="fig" rid="F1">Figure 1</xref> and Tooby and Cosmides, <xref ref-type="bibr" rid="B70">1992</xref>). For the hearing impaired, hearing aids can leverage the hard attentional competition between figure (speech) and background (noise), e.g., by using directional microphones, maybe even informed by physiological signals related to our listening intention (Mikkelsen et al., <xref ref-type="bibr" rid="B40">2015</xref>; Sch&#x000E4;fer et al., <xref ref-type="bibr" rid="B52">2018</xref>; Schroeer et al., <xref ref-type="bibr" rid="B54">2023</xref>, <xref ref-type="bibr" rid="B55">2024</xref>). Features such as exogenous cue weighting or dynamic processing modes, informed by evolutionary insights, may improve both safety and user experience in diverse environments (Carreti&#x000E9;, <xref ref-type="bibr" rid="B10">2014</xref>; Strauss et al., <xref ref-type="bibr" rid="B67">2024b</xref>; Edwards, <xref ref-type="bibr" rid="B18">2007</xref>). For example, the silent operation of electric cars poses risks by reducing the salience of auditory warning cues (Clendinning, <xref ref-type="bibr" rid="B13">2018</xref>). Artificially adding sound reintroduces a natural correspondence between auditory salience and threat but increases noise pollution (Hegewald et al., <xref ref-type="bibr" rid="B32">2020</xref>; Gilani and Mir, <xref ref-type="bibr" rid="B24">2021</xref>). An alternative is sonifying dynamic traffic information (see ETSI, <xref ref-type="bibr" rid="B19">2011</xref>) via bone-conduction devices, enabling the auditory system to use novel input sources while maintaining its evolved role as a 360&#x000B0; early-warning system (Strauss et al., <xref ref-type="bibr" rid="B64">2020</xref>; Olszanowski et al., <xref ref-type="bibr" rid="B43">2023</xref>). This approach highlights how hearing technologies can integrate evolutionarily honed mechanisms with modern demands. Numerical simulations of conceptual effortful listening models (e.g., see Schneider et al., <xref ref-type="bibr" rid="B53">2019</xref>; Strauss and Francis, <xref ref-type="bibr" rid="B66">2017</xref>) can further contribute to optimizing hearing aid designs, supporting users in navigating modern soundscapes.</p>
<p>Effortful listening arises from mismatches between auditory input and the brain&#x00027;s predictions, linked to increased free energy (Pichora-Fuller et al., <xref ref-type="bibr" rid="B48">2016</xref>; Strauss and Francis, <xref ref-type="bibr" rid="B66">2017</xref>). Predictive models show that internal demands related to a conflict of expectations across spatiotemporal scales and uncertainty drive listening effort, particularly in noisy environments. This framework connects listening effort to broader principles of brain function and allows for experimental exploration of multimodal integration (Calvert et al., <xref ref-type="bibr" rid="B9">2004</xref>; Schulte et al., <xref ref-type="bibr" rid="B57">2023</xref>) in effortful listening. Effortful listening informs the design of acoustic human-machine interfaces, particularly in noise-heavy environments like factories or vehicles (Neumann et al., <xref ref-type="bibr" rid="B42">2021</xref>; Damian et al., <xref ref-type="bibr" rid="B15">2015</xref>; Gonzalez-Trejo et al., <xref ref-type="bibr" rid="B26">2019</xref>). Neuroergonomic approaches (Parasuraman, <xref ref-type="bibr" rid="B44">2003</xref>) can minimize cognitive load by aligning design with the attentional system&#x00027;s evolutionary strengths, creating more intuitive and effective interfaces (Mehta and Parasuraman, <xref ref-type="bibr" rid="B39">2013</xref>). Thus, the computational implementation of the concepts presented in Section 2 (see also Bogacz, <xref ref-type="bibr" rid="B5">2017</xref>; Strauss et al., <xref ref-type="bibr" rid="B67">2024b</xref>) might support the optimization of neuroergonomic designs in medical, human-machine-interaction, and entertainment applications dealing with effortful listening.</p>
</sec>
<sec sec-type="conclusions" id="s4">
<title>4 Conclusions</title>
<p>We have considered the understanding of speech in noise within Anthropocene soundscapes from an evolutionary perspective. We propose that there have been, at most, marginal changes in auditory capacities over the last 200,000 years and essentially no changes in the last 2,000 years. However, the soundscape has changed drastically in just the last 200 years. Consequently, we propose that much of the effortful and unpleasant nature of extracting target speech and suppressing background noise stems from our auditory system not being adapted to modern acoustic environments. Using models that link effortful listening to attention, we examined the binary attentional competition between figure and background. We integrated these models into the free energy minimization framework to conceptualize attentional effort in listening. This evolutionary cognitive neuroscience approach and active inference model have implications for studying human audition and developing auditory technologies, including earbuds, hearing aids, immersive environments, and human-machine interaction systems.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s5">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/supplementary material, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec sec-type="author-contributions" id="s6">
<title>Author contributions</title>
<p>DS: Conceptualization, Formal analysis, Funding acquisition, Investigation, Supervision, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing, Methodology, Visualization. AF: Conceptualization, Formal analysis, Investigation, Methodology, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. ZS: Investigation, Methodology, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. ML: Investigation, Methodology, Writing &#x02013; review &#x00026; editing. FC-S: Conceptualization, Formal analysis, Investigation, Methodology, Writing &#x02013; review &#x00026; editing. SL: Conceptualization, Formal analysis, Investigation, Methodology, Supervision, Writing &#x02013; review &#x00026; editing.</p>
</sec>
<sec sec-type="funding-information" id="s7">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research and/or publication of this article. DS was partially supported by the German Federal Ministry of Education and Research (BMBF), Grant 13FH050KX1 &#x0201C;Deep Immersion Lab Saar&#x0201D; and the European Union (European Regional Development Fund, ERDF) and Saarland via the Center for Digital Neurotechnologies Saar (CDNS).</p>
</sec>
<ack><p>The authors would like to thank Steven A. Hackley, Clinical and Cognitive Neuroscience Laboratory, Department of Psychological Sciences, University of Missouri, USA for his valuable feedback on an earlier version of this article.</p>
</ack>
<sec sec-type="COI-statement" id="conf1">
<title>Conflict of interest</title>
<p>ML and SL were employed by Sonova AG.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="s8">
<title>Generative AI statement</title>
<p>The author(s) declare that Gen AI was used in the creation of this manuscript to generate the people in <xref ref-type="fig" rid="F2">Figure 2</xref> without copyright problems.</p>
</sec>
<sec sec-type="disclaimer" id="s9">
<title>Publisher&#x00027;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Aboitiz</surname> <given-names>F.</given-names></name></person-group> (<year>2018</year>). <article-title>A brain for speech. Evolutionary continuity in primate and human auditory-vocal processing</article-title>. <source>Front Neurosci</source>. <volume>12</volume>:<fpage>174</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2018.00174</pub-id><pub-id pub-id-type="pmid">29636657</pub-id></citation></ref>
<ref id="B2">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ackermann</surname> <given-names>H.</given-names></name> <name><surname>Hage</surname> <given-names>S. R.</given-names></name> <name><surname>Ziegler</surname> <given-names>W.</given-names></name></person-group> (<year>2014</year>). <article-title>Brain mechanisms of acoustic communication in humans and nonhuman primates: an evolutionary perspective</article-title>. <source>Behav. Brain Sci</source>. <volume>37</volume>, <fpage>529</fpage>&#x02013;<lpage>546</lpage>. <pub-id pub-id-type="doi">10.1017/S0140525X13003099</pub-id><pub-id pub-id-type="pmid">24827156</pub-id></citation></ref>
<ref id="B3">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ashe</surname> <given-names>A.</given-names></name> <name><surname>Colot</surname> <given-names>V.</given-names></name> <name><surname>Oldroyd</surname> <given-names>B. P.</given-names></name></person-group> (<year>2021</year>). <article-title>How does epigenetics influence the course of evolution?</article-title> <source>Philos. Trans. R. Soc. Lond. B. Biol. Sci</source>. <volume>376</volume>:<fpage>20200111</fpage>. <pub-id pub-id-type="doi">10.1098/rstb.2020.0111</pub-id><pub-id pub-id-type="pmid">33866814</pub-id></citation></ref>
<ref id="B4">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Boduroglu</surname> <given-names>A.</given-names></name> <name><surname>Shah</surname> <given-names>P.</given-names></name></person-group> (<year>2017</year>). <article-title>Cultural differences in attentional breadth and resolution</article-title>. <source>Cult. Brain</source> <volume>5</volume>, <fpage>169</fpage>&#x02013;<lpage>181</lpage>. <pub-id pub-id-type="doi">10.1007/s40167-017-0056-9</pub-id><pub-id pub-id-type="pmid">32895885</pub-id></citation></ref>
<ref id="B5">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bogacz</surname> <given-names>R.</given-names></name></person-group> (<year>2017</year>). <article-title>A tutorial on the free-energy framework for modelling perception and learning</article-title>. <source>J. Math. Psychol</source>. <volume>76</volume>, <fpage>198</fpage>&#x02013;<lpage>211</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmp.2015.11.003</pub-id><pub-id pub-id-type="pmid">28298703</pub-id></citation></ref>
<ref id="B6">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Boyd</surname> <given-names>R.</given-names></name> <name><surname>Richerson</surname> <given-names>P. J.</given-names></name> <name><surname>Henrich</surname> <given-names>J.</given-names></name></person-group> (<year>2013</year>). <source>The Cultural Evolution of Technology: Facts and Theories</source>. <publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>MIT Press</publisher-name>.</citation>
</ref>
<ref id="B7">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Bregman</surname> <given-names>A. S.</given-names></name></person-group> (<year>1990</year>). <source>Auditory Scene Analysis</source>. <publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>MIT Press</publisher-name>.</citation>
</ref>
<ref id="B8">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bruner</surname> <given-names>E.</given-names></name> <name><surname>Colom</surname> <given-names>R.</given-names></name></person-group> (<year>2022</year>). <article-title>Can a neandertal meditate? An evolutionary view of attention as a core component of general intelligence</article-title>. <source>Intelligence</source> <volume>93</volume>:<fpage>101668</fpage>. <pub-id pub-id-type="doi">10.1016/j.intell.2022.101668</pub-id></citation>
</ref>
<ref id="B9">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Calvert</surname> <given-names>G. A.</given-names></name> <name><surname>Spence</surname> <given-names>C.</given-names></name> <name><surname>Stein</surname> <given-names>B. E.</given-names></name></person-group> (<year>2004</year>). <source>The Handbook of Multisensory Processes</source>. <publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>MIT Press</publisher-name>.<pub-id pub-id-type="pmid">34448968</pub-id></citation></ref>
<ref id="B10">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Carreti&#x000E9;</surname> <given-names>L.</given-names></name></person-group> (<year>2014</year>). <article-title>Exogenous (automatic) attention to emotional stimuli: a review</article-title>. <source>Cogn. Affect. Behav. Neurosci</source>. <volume>14</volume>, <fpage>1228</fpage>&#x02013;<lpage>1258</lpage>. <pub-id pub-id-type="doi">10.3758/s13415-014-0270-2</pub-id><pub-id pub-id-type="pmid">24683062</pub-id></citation></ref>
<ref id="B11">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>Z.</given-names></name> <name><surname>Wiens</surname> <given-names>J. J.</given-names></name></person-group> (<year>2020</year>). <article-title>The origins of acoustic communication in vertebrates</article-title>. <source>Nat. Commun</source>. <volume>11</volume>:<fpage>369</fpage>. <pub-id pub-id-type="doi">10.1038/s41467-020-14356-3</pub-id><pub-id pub-id-type="pmid">31953401</pub-id></citation></ref>
<ref id="B12">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Clark</surname> <given-names>A.</given-names></name></person-group> (<year>2013</year>). <article-title>Whatever next? Predictive brains, situated agents, and the future of cognitive science</article-title>. <source>Behav. Brain Sci</source>. <volume>36</volume>, <fpage>181</fpage>&#x02013;<lpage>204</lpage>. <pub-id pub-id-type="doi">10.1017/S0140525X12000477</pub-id><pub-id pub-id-type="pmid">23663408</pub-id></citation></ref>
<ref id="B13">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Clendinning</surname> <given-names>E. A.</given-names></name></person-group> (<year>2018</year>). <article-title>Driving future sounds: imagination, identity and safety in electric vehicle noise design</article-title>. <source>Sound Stud</source>. <volume>4</volume>, <fpage>61</fpage>&#x02013;<lpage>76</lpage>. <pub-id pub-id-type="doi">10.1080/20551940.2018.1467664</pub-id></citation>
</ref>
<ref id="B14">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Conde-Valverde</surname> <given-names>M.</given-names></name> <name><surname>Mart&#x000ED;nez</surname> <given-names>I.</given-names></name> <name><surname>Quam</surname> <given-names>R. M.</given-names></name> <name><surname>Rosa</surname> <given-names>M.</given-names></name> <name><surname>Velez</surname> <given-names>A. D.</given-names></name> <name><surname>Lorenzo</surname> <given-names>C.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Neanderthals and homo sapiens had similar auditory and speech capacities</article-title>. <source>Nature Ecol. Evol</source>. <volume>5</volume>, <fpage>609</fpage>&#x02013;<lpage>615</lpage>. <pub-id pub-id-type="doi">10.1038/s41559-021-01391-6</pub-id><pub-id pub-id-type="pmid">33649543</pub-id></citation></ref>
<ref id="B15">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Damian</surname> <given-names>A.</given-names></name> <name><surname>Corona-Strauss</surname> <given-names>F. I.</given-names></name> <name><surname>Hannemann</surname> <given-names>R.</given-names></name> <name><surname>Strauss</surname> <given-names>D. J.</given-names></name></person-group> (<year>2015</year>). <article-title>&#x0201C;Towards the assessment of listening effort in real life situations: Mobile EEG recordings in a multimodal driving situation,&#x0201D;</article-title> in <source>2015 37th Annual International Conference of the IEEE Engineering in Medicine and Biology Society (EMBC)</source> (<publisher-loc>Milan</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>8123</fpage>&#x02013;<lpage>8126</lpage>.<pub-id pub-id-type="pmid">26738179</pub-id></citation></ref>
<ref id="B16">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Desimone</surname> <given-names>R.</given-names></name> <name><surname>Duncan</surname> <given-names>J.</given-names></name></person-group> (<year>1995</year>). <article-title>Neural mechanisms of selective visual attention</article-title>. <source>Annu. Rev. Neurosci</source>. <volume>18</volume>, <fpage>193</fpage>&#x02013;<lpage>222</lpage>. <pub-id pub-id-type="doi">10.1146/annurev.ne.18.030195.001205</pub-id><pub-id pub-id-type="pmid">7605061</pub-id></citation></ref>
<ref id="B17">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dichter</surname> <given-names>B. K.</given-names></name> <name><surname>Breshears</surname> <given-names>J. D.</given-names></name> <name><surname>Leonard</surname> <given-names>M. K.</given-names></name> <name><surname>Chang</surname> <given-names>E. F.</given-names></name></person-group> (<year>2018</year>). <article-title>The control of vocal pitch in human laryngeal motor cortex</article-title>. <source>Cell</source> <volume>174</volume>, <fpage>21</fpage>&#x02013;<lpage>31</lpage>.e9. <pub-id pub-id-type="doi">10.1016/j.cell.2018.05.016</pub-id><pub-id pub-id-type="pmid">29958109</pub-id></citation></ref>
<ref id="B18">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Edwards</surname> <given-names>E.</given-names></name></person-group> (<year>2007</year>). <article-title>The future of hearing aid technology</article-title>. <source>Trends Amplif</source>. <volume>11</volume>, <fpage>31</fpage>&#x02013;<lpage>45</lpage>. <pub-id pub-id-type="doi">10.1177/1084713806298004</pub-id><pub-id pub-id-type="pmid">17301336</pub-id></citation></ref>
<ref id="B19">
<citation citation-type="book"><person-group person-group-type="author"><collab>ETSI</collab></person-group> (<year>2011</year>). <source>Intelligent Transport Systems (ITS): Vehicular Communications; Basic Set of Applications; Local Dynamic Map (LDM) Rationale for and Guidance on Standardization. Tr 102 863 (v1.1.1)</source>. <publisher-loc>Sophia Antipolis</publisher-loc>: <publisher-name>European Telecommunications Standards Institute (ETSI)</publisher-name>.</citation>
</ref>
<ref id="B20">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Everett</surname> <given-names>D.</given-names></name></person-group> (<year>2017</year>). <source>How Language Began: The Story of Humanity&#x00027;s Greatest Invention</source>. <publisher-loc>New York, NY</publisher-loc>: <publisher-name>Liveright Publishing</publisher-name>.</citation>
</ref>
<ref id="B21">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Feldman</surname> <given-names>H.</given-names></name> <name><surname>Friston</surname> <given-names>K. J.</given-names></name></person-group> (<year>2010</year>). <article-title>Attention, uncertainty, and free-energy</article-title>. <source>Front. Hum. Neurosci</source>. <volume>4</volume>:<fpage>215</fpage>. <pub-id pub-id-type="doi">10.3389/fnhum.2010.00215</pub-id><pub-id pub-id-type="pmid">21160551</pub-id></citation></ref>
<ref id="B22">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Friston</surname> <given-names>K. J.</given-names></name></person-group> (<year>2010</year>). <article-title>The free-energy principle: a unified brain theory?</article-title> <source>Nat. Rev. Neurosci</source>. <volume>11</volume>, <fpage>127</fpage>&#x02013;<lpage>138</lpage>. <pub-id pub-id-type="doi">10.1038/nrn2787</pub-id><pub-id pub-id-type="pmid">20068583</pub-id></citation></ref>
<ref id="B23">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Gazzaley</surname> <given-names>A.</given-names></name> <name><surname>Rosen</surname> <given-names>L. D.</given-names></name></person-group> (<year>2016</year>). <source>The Distracted Mind: Ancient Brains in a High-Tech World</source>. <publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>MIT Press</publisher-name>.</citation>
</ref>
<ref id="B24">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gilani</surname> <given-names>T. A.</given-names></name> <name><surname>Mir</surname> <given-names>M. S.</given-names></name></person-group> (<year>2021</year>). <article-title>A study on the assessment of traffic noise induced annoyance and awareness levels about the potential health effects among residents living around a noise-sensitive area</article-title>. <source>Environ. Sci. Pollut. Res. Int</source>. <volume>28</volume>, <fpage>63045</fpage>&#x02013;<lpage>63064</lpage>. <pub-id pub-id-type="doi">10.1007/s11356-021-15208-3</pub-id><pub-id pub-id-type="pmid">34218377</pub-id></citation></ref>
<ref id="B25">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gintis</surname> <given-names>H.</given-names></name></person-group> (<year>2011</year>). <article-title>Gene-culture coevolution and the nature of human sociality</article-title>. <source>Philos. Trans. R. Soc. Lond. B. Biol. Sci</source>. <volume>366</volume>, <fpage>878</fpage>&#x02013;<lpage>888</lpage>. <pub-id pub-id-type="doi">10.1098/rstb.2010.0310</pub-id><pub-id pub-id-type="pmid">21320901</pub-id></citation></ref>
<ref id="B26">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gonzalez-Trejo</surname> <given-names>E.</given-names></name> <name><surname>M&#x000F6;gele</surname> <given-names>H.</given-names></name> <name><surname>Pfleger</surname> <given-names>N.</given-names></name> <name><surname>Hannemann</surname> <given-names>R.</given-names></name> <name><surname>Strauss</surname> <given-names>D. J.</given-names></name></person-group> (<year>2019</year>). <article-title>Electroencephalographic phase-amplitude coupling in simulated driving with varying modality-specific attentional demand</article-title>. <source>IEEE Trans. Human-Mach. Syst</source>. <volume>49</volume>, <fpage>589</fpage>&#x02013;<lpage>598</lpage>. <pub-id pub-id-type="doi">10.1109/THMS.2019.2931011</pub-id></citation>
</ref>
<ref id="B27">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Grossberg</surname> <given-names>S.</given-names></name></person-group> (<year>2005</year>). <article-title>&#x0201C;Linking attention to learning, expectation, competition, and consciousness,&#x0201D;</article-title> in <source>Neurobiology of Attention</source>, eds. L. Itti, and J. Tsotsos (Burlington: Academic Press), <fpage>652</fpage>&#x02013;<lpage>662</lpage>.</citation>
</ref>
<ref id="B28">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Grothe</surname> <given-names>B.</given-names></name> <name><surname>Carr</surname> <given-names>C. E.</given-names></name> <name><surname>Casseday</surname> <given-names>J. H.</given-names></name> <name><surname>Fritzsch</surname> <given-names>B.</given-names></name> <name><surname>K&#x000F6;ppl</surname> <given-names>C.</given-names></name></person-group> (<year>2004</year>). <article-title>&#x0201C;The evolution of central pathways and their neural processing patterns,&#x0201D;</article-title> in <source>Evolution of the Vertebrate Auditory System</source>, eds. G. A. Manley, R. R. Fay, and A. N. Popper (New York, NY: Springer), <fpage>289</fpage>&#x02013;<lpage>359</lpage>.</citation>
</ref>
<ref id="B29">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Habib</surname> <given-names>L.</given-names></name> <name><surname>Bayne</surname> <given-names>E. M.</given-names></name> <name><surname>Boutin</surname> <given-names>S.</given-names></name></person-group> (<year>2007</year>). <article-title>Chronic industrial noise affects pairing success and age structure of ovenbirds seiurus aurocapilla</article-title>. <source>J. Appl. Ecol</source>. <volume>44</volume>, <fpage>176</fpage>&#x02013;<lpage>184</lpage>. <pub-id pub-id-type="doi">10.1111/j.1365-2664.2006.01234.x</pub-id></citation>
</ref>
<ref id="B30">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hackley</surname> <given-names>S. A.</given-names></name></person-group> (<year>2015</year>). <article-title>Evidence for a vestigial pinna-orienting system in humans</article-title>. <source>Psychophysiology</source> <volume>52</volume>, <fpage>1263</fpage>&#x02013;<lpage>1270</lpage>. <pub-id pub-id-type="doi">10.1111/psyp.12501</pub-id><pub-id pub-id-type="pmid">26211937</pub-id></citation></ref>
<ref id="B31">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hawks</surname> <given-names>J.</given-names></name> <name><surname>Wang</surname> <given-names>E. T.</given-names></name> <name><surname>Cochran</surname> <given-names>G. M.</given-names></name> <name><surname>Harpending</surname> <given-names>H. C.</given-names></name> <name><surname>Moyzis</surname> <given-names>R. K.</given-names></name></person-group> (<year>2007</year>). <article-title>Recent acceleration of human adaptive evolution</article-title>. <source>Proc. Natl. Acad. Sci</source>. <volume>104</volume>, <fpage>20753</fpage>&#x02013;<lpage>20758</lpage>. <pub-id pub-id-type="doi">10.1073/pnas.0707650104</pub-id><pub-id pub-id-type="pmid">18087044</pub-id></citation></ref>
<ref id="B32">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hegewald</surname> <given-names>J.</given-names></name> <name><surname>Schubert</surname> <given-names>M.</given-names></name> <name><surname>Freiberg</surname> <given-names>A.</given-names></name> <name><surname>Romero Starke</surname> <given-names>K.</given-names></name> <name><surname>Augustin</surname> <given-names>F.</given-names></name> <name><surname>Riedel-Heller</surname> <given-names>S. G.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>Traffic noise and mental health: a systematic review and meta-analysis</article-title>. <source>Int. J. Environ. Res. Public Health</source> <volume>17</volume>:<fpage>6175</fpage>. <pub-id pub-id-type="doi">10.3390/ijerph17176175</pub-id><pub-id pub-id-type="pmid">32854453</pub-id></citation></ref>
<ref id="B33">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jigo</surname> <given-names>M.</given-names></name> <name><surname>Heeger</surname> <given-names>D. J.</given-names></name> <name><surname>Carrasco</surname> <given-names>M.</given-names></name></person-group> (<year>2021</year>). <article-title>An image-computable model of how endogenous and exogenous attention differentially alter visual perception</article-title>. <source>Proc. Natl. Acad. Sci</source>. <volume>17</volume>:<fpage>e2106436118</fpage>. <pub-id pub-id-type="doi">10.1101/2021.01.26.428173</pub-id><pub-id pub-id-type="pmid">34389680</pub-id></citation></ref>
<ref id="B34">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jurkat</surname> <given-names>S.</given-names></name> <name><surname>Koster</surname> <given-names>M.</given-names></name> <name><surname>Yovsi</surname> <given-names>R.</given-names></name></person-group> (<year>2020</year>). <article-title>The development of context-sensitive attention across cultures: the impact of stimulus familiarity</article-title>. <source>Front. Psychol</source>. <volume>11</volume>:<fpage>1526</fpage>. <pub-id pub-id-type="doi">10.3389/fpsyg.2020.01526</pub-id><pub-id pub-id-type="pmid">32760322</pub-id></citation></ref>
<ref id="B35">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lev-Ari</surname> <given-names>T.</given-names></name> <name><surname>Beeri</surname> <given-names>H.</given-names></name> <name><surname>Gutfreund</surname> <given-names>Y.</given-names></name></person-group> (<year>2022</year>). <article-title>The ecological view of selective attention</article-title>. <source>Front. Integr. Neurosci</source>. <volume>16</volume>:<fpage>856207</fpage>. <pub-id pub-id-type="doi">10.3389/fnint.2022.856207</pub-id><pub-id pub-id-type="pmid">35391754</pub-id></citation></ref>
<ref id="B36">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lorenzi</surname> <given-names>C.</given-names></name> <name><surname>Apoux</surname> <given-names>F.</given-names></name> <name><surname>Grinfeder</surname> <given-names>E.</given-names></name> <name><surname>Krause</surname> <given-names>B.</given-names></name> <name><surname>Miller-Viacava</surname> <given-names>N.</given-names></name> <name><surname>Sueur</surname> <given-names>J.</given-names></name></person-group> (<year>2023</year>). <article-title>Human auditory ecology: extending hearing research to the perception of natural soundscapes by humans in rapidly changing environment</article-title>. <source>Trends Hear</source>. <volume>17</volume>:<fpage>23312165231212032</fpage>. <pub-id pub-id-type="doi">10.1177/23312165231212032</pub-id><pub-id pub-id-type="pmid">37981813</pub-id></citation></ref>
<ref id="B37">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Manley</surname> <given-names>G. A.</given-names></name></person-group> (<year>2017</year>). <article-title>Comparative auditory neuroscience: understanding the evolution and function of ears</article-title>. <source>J. Assoc. Res. Otolaryngol</source>. <volume>18</volume>, <fpage>1</fpage>&#x02013;<lpage>24</lpage>. <pub-id pub-id-type="doi">10.1007/s10162-016-0579-3</pub-id><pub-id pub-id-type="pmid">27539715</pub-id></citation></ref>
<ref id="B38">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Marr</surname> <given-names>D.</given-names></name></person-group> (<year>1982</year>). <source>Vision</source>. <publisher-loc>San Francisco</publisher-loc>: <publisher-name>Freeman</publisher-name>.</citation>
</ref>
<ref id="B39">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mehta</surname> <given-names>R. K.</given-names></name> <name><surname>Parasuraman</surname> <given-names>R.</given-names></name></person-group> (<year>2013</year>). <article-title>Neuroergonomics: a review of application to physical and cognitive work</article-title>. <source>Front. Hum. Neurosci</source>. <volume>7</volume>:<fpage>889</fpage>. <pub-id pub-id-type="doi">10.3389/fnhum.2013.00889</pub-id><pub-id pub-id-type="pmid">24391575</pub-id></citation></ref>
<ref id="B40">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mikkelsen</surname> <given-names>K. B.</given-names></name> <name><surname>Kappel</surname> <given-names>S. L.</given-names></name> <name><surname>Mandic</surname> <given-names>D. P.</given-names></name> <name><surname>Kidmose</surname> <given-names>P.</given-names></name></person-group> (<year>2015</year>). <article-title>EEG recorded from the ear: characterizing the ear-EEG method</article-title>. <source>Front. Neurosci</source>. <volume>9</volume>:<fpage>109076</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2015.00438</pub-id><pub-id pub-id-type="pmid">26635514</pub-id></citation></ref>
<ref id="B41">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>M&#x000FC;ller</surname> <given-names>H. J.</given-names></name> <name><surname>Rabbitt</surname> <given-names>P. M.</given-names></name></person-group> (<year>1989</year>). <article-title>Reflexive and voluntary orienting of visual attention: time course of activation and resistance to interruption</article-title>. <source>J. Exp. Psychol. Hum. Percept. Perform</source>. <volume>15</volume>:<fpage>315</fpage>&#x02013;<lpage>330</lpage>. <pub-id pub-id-type="doi">10.1037//0096-1523.15.2.315</pub-id><pub-id pub-id-type="pmid">2525601</pub-id></citation></ref>
<ref id="B42">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Neumann</surname> <given-names>W. P.</given-names></name> <name><surname>Winkelhaus</surname> <given-names>S.</given-names></name> <name><surname>Grosse</surname> <given-names>E. H.</given-names></name> <name><surname>Glock</surname> <given-names>C. H.</given-names></name></person-group> (<year>2021</year>). <article-title>Industry 4.0 and the human factor &#x02013; a systems framework and analysis methodology for successful development</article-title>. <source>Int. J. Prod. Econ</source>. <volume>233</volume>:<fpage>107992</fpage>. <pub-id pub-id-type="doi">10.1016/j.ijpe.2020.107992</pub-id></citation>
</ref>
<ref id="B43">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Olszanowski</surname> <given-names>M.</given-names></name> <name><surname>Frankowska</surname> <given-names>N.</given-names></name> <name><surname>Tolopilo</surname> <given-names>A.</given-names></name></person-group> (<year>2023</year>). <article-title>&#x0201D;Rear bias&#x0201D; in spatial auditory perception: attentional and affective vigilance to sounds occurring outside the visual field</article-title>. <source>Psychophysiology</source> <volume>60</volume>:<fpage>e14377</fpage>. <pub-id pub-id-type="doi">10.1111/psyp.14377</pub-id><pub-id pub-id-type="pmid">37357967</pub-id></citation></ref>
<ref id="B44">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Parasuraman</surname> <given-names>R.</given-names></name></person-group> (<year>2003</year>). <article-title>Neuroergonomics: research and practice</article-title>. <source>Theoret. Issues in Ergon. Sci</source>. <volume>4</volume>, <fpage>5</fpage>&#x02013;<lpage>20</lpage>. <pub-id pub-id-type="doi">10.1080/14639220210199753</pub-id></citation>
</ref>
<ref id="B45">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Parr</surname> <given-names>T.</given-names></name> <name><surname>Pezzulo</surname> <given-names>G.</given-names></name> <name><surname>Friston</surname> <given-names>K. J.</given-names></name></person-group> (<year>2022</year>). <source>Active Inference: the Free Energy Principle in Mind, Brain, and Behavior</source>. <publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>MIT Press</publisher-name>.</citation>
</ref>
<ref id="B46">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Parvizi-Wayne</surname> <given-names>D.</given-names></name></person-group> (<year>2024</year>). <article-title>How preference enslaves attention: calling into question the endogenous/exogenous distinction from an active inference perspective</article-title>. <source>Phenomenol. Cogn. Sci.</source> <pub-id pub-id-type="doi">10.1007/s11097-024-10028-5</pub-id></citation>
</ref>
<ref id="B47">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Patel</surname> <given-names>A. D.</given-names></name></person-group> (<year>2021</year>). <article-title>Vocal learning as a preadaptation for the evolution of human beat perception and synchronization</article-title>. <source>Philos. Trans. R. Soc. Lond. B. Biol. Sci</source>. <volume>376</volume>:<fpage>20200326</fpage>. <pub-id pub-id-type="doi">10.1098/rstb.2020.0326</pub-id><pub-id pub-id-type="pmid">34420384</pub-id></citation></ref>
<ref id="B48">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Pichora-Fuller</surname> <given-names>M. K.</given-names></name> <name><surname>Kramer</surname> <given-names>S. E.</given-names></name> <name><surname>Eckert</surname> <given-names>M. A.</given-names></name> <name><surname>Edwards</surname> <given-names>B.</given-names></name> <name><surname>Hornsby</surname> <given-names>B. W.</given-names></name> <name><surname>Humes</surname> <given-names>L. E.</given-names></name> <etal/></person-group>. (<year>2016</year>). <article-title>Hearing impairment and cognitive energy: the framework for understanding effortful listening (FUEL)</article-title>. <source>Ear Hear</source>. <volume>37</volume>, <fpage>5S</fpage>&#x02013;<lpage>27S</lpage>. <pub-id pub-id-type="doi">10.1097/AUD.0000000000000312</pub-id><pub-id pub-id-type="pmid">27355771</pub-id></citation></ref>
<ref id="B49">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ren</surname> <given-names>Y.</given-names></name> <name><surname>Zhang</surname> <given-names>Y.</given-names></name> <name><surname>Hou</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>J.</given-names></name> <name><surname>Bi</surname> <given-names>J.</given-names></name> <name><surname>Yang</surname> <given-names>W.</given-names></name></person-group> (<year>2021</year>). <article-title>Exogenous bimodal cues attenuate age-related audiovisual integration</article-title>. <source>Iperception</source> <volume>12</volume>:<fpage>20416695211020768</fpage>. <pub-id pub-id-type="doi">10.1177/20416695211020768</pub-id><pub-id pub-id-type="pmid">34104386</pub-id></citation></ref>
<ref id="B50">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rosenblum</surname> <given-names>L. D.</given-names></name></person-group> (<year>2008</year>). <article-title>Speech perception as a multimodal phenomenon</article-title>. <source>Curr. Dir. Psychol. Sci</source>. <volume>17</volume>, <fpage>405</fpage>&#x02013;<lpage>409</lpage>. <pub-id pub-id-type="doi">10.1111/j.1467-8721.2008.00615.x</pub-id><pub-id pub-id-type="pmid">23914077</pub-id></citation></ref>
<ref id="B51">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sarter</surname> <given-names>M.</given-names></name> <name><surname>Gehring</surname> <given-names>W.</given-names></name> <name><surname>Kozak</surname> <given-names>R.</given-names></name></person-group> (<year>2006</year>). <article-title>More attention must be paid: the neurobiology of attentional effort</article-title>. <source>Brain Res. Rev</source>. <volume>51</volume>, <fpage>145</fpage>&#x02013;<lpage>160</lpage>. <pub-id pub-id-type="doi">10.1016/j.brainresrev.2005.11.002</pub-id><pub-id pub-id-type="pmid">16530842</pub-id></citation></ref>
<ref id="B52">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sch&#x000E4;fer</surname> <given-names>P. J.</given-names></name> <name><surname>Corona-Strauss</surname> <given-names>F. I.</given-names></name> <name><surname>Hannemann</surname> <given-names>R.</given-names></name> <name><surname>Hillyard</surname> <given-names>S. A.</given-names></name> <name><surname>Strauss</surname> <given-names>D. J.</given-names></name></person-group> (<year>2018</year>). <article-title>Testing the limits of the stimulus reconstruction approach: auditory attention decoding in a four-speaker free field environment</article-title>. <source>Trends Hear</source>. <volume>22</volume>:<fpage>233121651881660</fpage>. <pub-id pub-id-type="doi">10.1177/2331216518816600</pub-id></citation>
</ref>
<ref id="B53">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Schneider</surname> <given-names>E. N.</given-names></name> <name><surname>Bernarding</surname> <given-names>C.</given-names></name> <name><surname>Francis</surname> <given-names>A. L.</given-names></name> <name><surname>Hornsby</surname> <given-names>B. W. Y.</given-names></name> <name><surname>Strauss</surname> <given-names>D. J.</given-names></name></person-group> (<year>2019</year>). <article-title>&#x0201C;A quantitative model of listening related fatigue. Neural Engineering (NER),&#x0201D;</article-title> in <source>9th International IEEE/EMBS Conference</source> (<publisher-loc>San Francisco, CA</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>619</fpage>&#x02013;<lpage>622</lpage>.</citation>
</ref>
<ref id="B54">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schroeer</surname> <given-names>A.</given-names></name> <name><surname>Andersen</surname> <given-names>M. R.</given-names></name> <name><surname>Rank</surname> <given-names>M. L.</given-names></name> <name><surname>Hannemann</surname> <given-names>R.</given-names></name> <name><surname>Petersen</surname> <given-names>E. B.</given-names></name> <name><surname>R</surname> <given-names>F. M.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Assessment of vestigial auriculomotor activity to acoustic stimuli using electrodes in and around the ear</article-title>. <source>Trends Hear</source>. <volume>27</volume>:<fpage>23312165231200158</fpage>. <pub-id pub-id-type="doi">10.1177/23312165231200158</pub-id><pub-id pub-id-type="pmid">37830146</pub-id></citation></ref>
<ref id="B55">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schroeer</surname> <given-names>A.</given-names></name> <name><surname>Corona-Strauss</surname> <given-names>F. I.</given-names></name> <name><surname>Hannemann</surname> <given-names>R.</given-names></name> <name><surname>Hackley</surname> <given-names>S. A.</given-names></name> <name><surname>Strauss</surname> <given-names>D. J.</given-names></name></person-group> (<year>2024</year>). <article-title>Electromyographic correlates of effortful listening in the vestigial auriculomotor system</article-title>. <source>Front. Neurosci</source>. <volume>18</volume>:<fpage>1462507</fpage>. <pub-id pub-id-type="doi">10.3389/fnins.2024.1462507</pub-id><pub-id pub-id-type="pmid">39959575</pub-id></citation></ref>
<ref id="B56">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sch&#x000FC;lke</surname> <given-names>O.</given-names></name> <name><surname>Dumdey</surname> <given-names>N.</given-names></name> <name><surname>Ostner</surname> <given-names>J.</given-names></name></person-group> (<year>2020</year>). <article-title>Selective attention for affiliative and agonistic interactions of dominants and close affiliates in macaques</article-title>. <source>Sci. Rep</source>. <volume>10</volume>:<fpage>5962</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-020-62772-8</pub-id><pub-id pub-id-type="pmid">32249792</pub-id></citation></ref>
<ref id="B57">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schulte</surname> <given-names>A.</given-names></name> <name><surname>Marozeau</surname> <given-names>J.</given-names></name> <name><surname>Ruhe</surname> <given-names>A.</given-names></name> <name><surname>B&#x000FC;chner</surname> <given-names>A.</given-names></name> <name><surname>Kral</surname> <given-names>A.</given-names></name> <name><surname>Innes-Brown</surname> <given-names>H.</given-names></name></person-group> (<year>2023</year>). <article-title>Improved speech intelligibility in the presence of congruent vibrotactile speech input</article-title>. <source>Sci. Rep</source>. <volume>13</volume>:<fpage>22657</fpage>. <pub-id pub-id-type="doi">10.1038/s41598-023-48893-w</pub-id><pub-id pub-id-type="pmid">38114599</pub-id></citation></ref>
<ref id="B58">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shinn-Cunningham</surname> <given-names>B. G.</given-names></name></person-group> (<year>2008</year>). <article-title>Object-based auditory and visual attention</article-title>. <source>Trends Cogn. Sci</source>. <volume>12</volume>, <fpage>182</fpage>&#x02013;<lpage>186</lpage>. <pub-id pub-id-type="doi">10.1016/j.tics.2008.02.003</pub-id><pub-id pub-id-type="pmid">18396091</pub-id></citation></ref>
<ref id="B59">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Slabbekoorn</surname> <given-names>H.</given-names></name></person-group> (<year>2018</year>). <article-title>Soundscape ecology of the anthropocene</article-title>. <source>Acoust. Today</source> <volume>14</volume>, <fpage>42</fpage>&#x02013;<lpage>49</lpage>.</citation>
</ref>
<ref id="B60">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Star</surname> <given-names>B.</given-names></name> <name><surname>Spencer</surname> <given-names>H. G.</given-names></name></person-group> (<year>2013</year>). <article-title>Effects of genetic drift and gene flow on the selective maintenance of genetic variation</article-title>. <source>Genetics</source> <volume>194</volume>, <fpage>235</fpage>&#x02013;<lpage>244</lpage>. <pub-id pub-id-type="doi">10.1534/genetics.113.149781</pub-id><pub-id pub-id-type="pmid">23457235</pub-id></citation></ref>
<ref id="B61">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Sterbing-D&#x00027;Angelo</surname> <given-names>S. J.</given-names></name></person-group> (<year>2009</year>). <source>Encyclopedia of Neuroscience, chapter Evolution of the Auditory System</source>. <publisher-loc>Berlin, Heidelberg</publisher-loc>: <publisher-name>Springer</publisher-name>, <fpage>1286</fpage>&#x02013;<lpage>1288</lpage>.</citation>
</ref>
<ref id="B62">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Stevens</surname> <given-names>M.</given-names></name></person-group> (<year>2013</year>). <source>Sensory Ecology, Behavior, and Evolution</source>. <publisher-loc>Oxford</publisher-loc>: <publisher-name>Oxford University Press</publisher-name>.</citation>
</ref>
<ref id="B63">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Strauss</surname> <given-names>D. J.</given-names></name> <name><surname>Corona-Strauss</surname> <given-names>F. I.</given-names></name> <name><surname>Mai</surname> <given-names>A.</given-names></name> <name><surname>Hillyard</surname> <given-names>S. A.</given-names></name></person-group> (<year>2024a</year>). <article-title>Fifty years after: The n1 effect travels down to the brainstem</article-title>. <source>BioRxiv</source>. <pub-id pub-id-type="doi">10.1101/2024.02.23.581747</pub-id></citation>
</ref>
<ref id="B64">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Strauss</surname> <given-names>D. J.</given-names></name> <name><surname>Corona-Strauss</surname> <given-names>F. I.</given-names></name> <name><surname>Schroeer</surname> <given-names>A.</given-names></name> <name><surname>Flotho</surname> <given-names>P.</given-names></name> <name><surname>Hannemann</surname> <given-names>R.</given-names></name> <name><surname>Hackley</surname> <given-names>S. A.</given-names></name></person-group> (<year>2020</year>). <article-title>Vestigial auriculomotor activity indicates the direction of auditory attention in humans</article-title>. <source>Elife</source> <volume>9</volume>:<fpage>e54536</fpage>. <pub-id pub-id-type="doi">10.7554/eLife.54536</pub-id><pub-id pub-id-type="pmid">32618268</pub-id></citation></ref>
<ref id="B65">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Strauss</surname> <given-names>D. J.</given-names></name> <name><surname>Corona-Strauss</surname> <given-names>F. I.</given-names></name> <name><surname>Trenado</surname> <given-names>C.</given-names></name> <name><surname>Bernarding</surname> <given-names>C.</given-names></name> <name><surname>Reith</surname> <given-names>W.</given-names></name> <name><surname>Latzel</surname> <given-names>M.</given-names></name> <etal/></person-group>. (<year>2010</year>). <article-title>Electrophysiological correlates of listening effort: Neurodynamical modeling and measurement</article-title>. <source>Cogn. Neurodyn</source>. <volume>4</volume>, <fpage>119</fpage>&#x02013;<lpage>131</lpage>. <pub-id pub-id-type="doi">10.1007/s11571-010-9111-3</pub-id><pub-id pub-id-type="pmid">21629585</pub-id></citation></ref>
<ref id="B66">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Strauss</surname> <given-names>D. J.</given-names></name> <name><surname>Francis</surname> <given-names>A. J.</given-names></name></person-group> (<year>2017</year>). <article-title>Toward taxonomic model of attention in effortful listening</article-title>. <source>Cogn. Affect. Behav. Neurosci</source>. <volume>17</volume>, <fpage>809</fpage>&#x02013;<lpage>825</lpage>. <pub-id pub-id-type="doi">10.3758/s13415-017-0513-0</pub-id><pub-id pub-id-type="pmid">28567568</pub-id></citation></ref>
<ref id="B67">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Strauss</surname> <given-names>D. J.</given-names></name> <name><surname>Francis</surname> <given-names>A. L.</given-names></name> <name><surname>Vibell</surname> <given-names>J.</given-names></name> <name><surname>Corona-Strauss</surname> <given-names>F. I.</given-names></name></person-group> (<year>2024b</year>). <article-title>The role of attention in immersion: the two-competitor model</article-title>. <source>Brain Res. Bull</source>. <volume>210</volume>:<fpage>110923</fpage>. <pub-id pub-id-type="doi">10.1016/j.brainresbull.2024.110923</pub-id><pub-id pub-id-type="pmid">38462137</pub-id></citation></ref>
<ref id="B68">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Swaddle</surname> <given-names>J. P.</given-names></name> <name><surname>Francis</surname> <given-names>C.</given-names></name> <name><surname>Barber</surname> <given-names>J.</given-names></name> <name><surname>Cooper</surname> <given-names>C.</given-names></name> <name><surname>Kyba</surname> <given-names>C.</given-names></name> <name><surname>Dominoni</surname> <given-names>D.</given-names></name> <etal/></person-group>. (<year>2015</year>). <article-title>A framework to assess evolutionary responses to anthropogenic light and sound</article-title>. <source>Trends Ecol. Evol</source>. <volume>30</volume>, <fpage>550</fpage>&#x02013;<lpage>560</lpage>. <pub-id pub-id-type="doi">10.1016/j.tree.2015.06.009</pub-id><pub-id pub-id-type="pmid">26169593</pub-id></citation></ref>
<ref id="B69">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Swedell</surname> <given-names>L.</given-names></name> <name><surname>Plummer</surname> <given-names>T.</given-names></name></person-group> (<year>2019</year>). <article-title>Social evolution in plio-pleistocene hominins: insights from hamadryas baboons and paleoecology</article-title>. <source>J. Hum. Evol</source>. <volume>137</volume>:<fpage>102667</fpage>. <pub-id pub-id-type="doi">10.1016/j.jhevol.2019.102667</pub-id><pub-id pub-id-type="pmid">31629289</pub-id></citation></ref>
<ref id="B70">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tooby</surname> <given-names>J.</given-names></name> <name><surname>Cosmides</surname> <given-names>L.</given-names></name></person-group> (<year>1992</year>). <article-title>&#x0201C;The psychological foundations of culture,&#x0201D;</article-title> in <source>The Adapted Mind: Evolutionary Psychology and the Generation of Culture, Vol. 19</source>, eds. J. H. Barkow, L. Cosmides, and J. Tooby (New Yord, NY: Oxford University Press), <fpage>1</fpage>&#x02013;<lpage>136</lpage>.</citation>
</ref>
<ref id="B71">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Trenado</surname> <given-names>C.</given-names></name> <name><surname>Haab</surname> <given-names>L.</given-names></name> <name><surname>Strauss</surname> <given-names>D. J.</given-names></name></person-group> (<year>2009</year>). <article-title>Corticothalamic feedback dynamics for neural correlates of auditory selective attention</article-title>. <source>IEEE Trans. Neural Syst. Rehabil. Eng</source>. <volume>17</volume>, <fpage>46</fpage>&#x02013;<lpage>52</lpage>. <pub-id pub-id-type="doi">10.1109/TNSRE.2008.2010469</pub-id><pub-id pub-id-type="pmid">19211323</pub-id></citation></ref>
</ref-list>
</back>
</article>