<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Psychol.</journal-id>
<journal-title>Frontiers in Psychology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Psychol.</abbrev-journal-title>
<issn pub-type="epub">1664-1078</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpsyg.2024.1383904</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Psychology</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>From first encounters to longitudinal exposure: a repeated exposure-test paradigm for monitoring speech adaptation</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name><surname>Xie</surname> <given-names>Xin</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2386411/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Kurumada</surname> <given-names>Chigusa</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2698537/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Department of Language Science, University of California, Irvine</institution>, <addr-line>Irvine, CA</addr-line>, <country>United States</country></aff>
<aff id="aff2"><sup>2</sup><institution>Department of Brain and Cognitive Sciences, University of Rochester</institution>, <addr-line>Rochester, NY</addr-line>, <country>United States</country></aff>
<author-notes>
<fn fn-type="edited-by" id="fn0003">
<p>Edited by: Juhani J&#x00E4;rvikivi, University of Alberta, Canada</p>
</fn>
<fn fn-type="edited-by" id="fn0004">
<p>Reviewed by: Katie Von Holzen, Universit&#x00E9; Paris Cit&#x00E9;, France</p>
<p>Yevgeniy Melguy, Basque Center on Cognition, Brain and Language, Spain</p>
</fn>
<corresp id="c001">&#x002A;Correspondence: Xin Xie, <email>xxie14@uci.edu</email></corresp>
</author-notes>
<pub-date pub-type="epub">
<day>30</day>
<month>05</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>15</volume>
<elocation-id>1383904</elocation-id>
<history>
<date date-type="received">
<day>08</day>
<month>02</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>08</day>
<month>05</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2024 Xie and Kurumada.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Xie and Kurumada</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>Perceptual difficulty with an unfamiliar accent can dissipate within short time scales (e.g., within minutes), reflecting rapid adaptation effects. At the same time, long-term familiarity with an accent is also known to yield stable perceptual benefits. However, whether the long-term effects reflect sustained, cumulative progression from shorter-term adaptation remains unknown. To fill this gap, we developed a web-based, repeated exposure-test paradigm. In this paradigm, short test blocks alternate with exposure blocks, and this exposure-test sequence is repeated multiple times. This design allows for the testing of adaptive speech perception both (a) within the first moments of encountering an unfamiliar accent and (b) over longer time scales such as days and weeks. In addition, we used a Bayesian ideal observer approach to select natural speech stimuli that increase the statistical power to detect adaptation. The current report presents results from a first application of this paradigm, investigating changes in the recognition accuracy of Mandarin-accented speech by native English listeners over five sessions spanning 3&#x2009;weeks. We found that the recognition of an accent feature (a syllable-final /d/, as in <italic>feed</italic>, sounding/t/-like) improved steadily over the three-week period. Unexpectedly, however, the improvement was seen with or without exposure to the accent. We discuss possible reasons for this result and implications for conducting future longitudinal studies with repeated exposure and testing.</p>
</abstract>
<kwd-group>
<kwd>nonnative accent</kwd>
<kwd>adaptive speech perception</kwd>
<kwd>repeated exposure-tests</kwd>
<kwd>longitudinal testing</kwd>
<kwd>online perceptual experiment</kwd>
</kwd-group>
<counts>
<fig-count count="5"/>
<table-count count="0"/>
<equation-count count="0"/>
<ref-count count="44"/>
<page-count count="10"/>
<word-count count="7150"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Psychology of Language</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec1">
<label>1</label>
<title>Introduction</title>
<p>How listeners navigate the substantial amount of cross-talker variability is a central question in speech perception. The &#x201C;same&#x201D; phonological category or word is produced with distinct acoustic-phonetic properties across talkers with different characteristics (e.g., height, gender, accent). This variability is known to make the recognition of unfamiliar talkers or accents difficult (<xref ref-type="bibr" rid="ref1">Adank and Janse, 2010</xref>; <xref ref-type="bibr" rid="ref26">Porretta et al., 2016</xref>). These difficulties can, however, dissipate as listeners adapt to the current input (<xref ref-type="bibr" rid="ref10">Bradlow and Bent, 2008</xref>; <xref ref-type="bibr" rid="ref28">Tzeng et al., 2016</xref>; <xref ref-type="bibr" rid="ref3">Baese-Berk et al., 2020</xref>; <xref ref-type="bibr" rid="ref36">Xie et al., 2021</xref>). For example, native listeners of English become significantly faster and more accurate in responding to Spanish-or Mandarin-accented speech within as few as 18 sentence-length utterances (<xref ref-type="bibr" rid="ref12">Clarke and Garrett, 2004</xref>; <xref ref-type="bibr" rid="ref39">Xie et al., 2018</xref>).</p>
<p>Among the speech variants used to study adaptive perception, nonnative accents have several unique properties. Most prominent are the complex ways in which they deviate from the native variants, both at the segmental and suprasegmental levels. Unlike other types of acoustically degraded or noisy speech, accented speech is difficult to understand primarily because it alters how acoustic cues map onto speech categories such as phonemes and words. In some cases, one category is phonetically confusable with another (e.g., a voiced stop consonant like the [d] in &#x201C;seed&#x201D; is often devoiced in a word-final position that sounds more like the [t] in &#x201C;seat&#x201D; in German and Dutch accented English, <xref ref-type="bibr" rid="ref16">Eisner et al., 2013</xref>); in others, categories are merged, substituted, or omitted (e.g., the English /&#x03B8;/ is substituted by different categories across accents <xref ref-type="bibr" rid="ref20">Hanul&#x00ED;kov&#x00E1; and Weber, 2012</xref>, for a review see <xref ref-type="bibr" rid="ref6">Bent and Baese-Berk, 2021</xref>). These variations can lead to lexical ambiguity and uncertainty, often resulting in slower and less accurate recognition.</p>
<p>While these variations may be idiosyncratic, they are far from random. Talkers from similar native language (L1) backgrounds tend to share common accent features, influenced by L1 phonology and its difference from the nonnative (L2) phonology (<xref ref-type="bibr" rid="ref18">Flege et al., 1992</xref>; <xref ref-type="bibr" rid="ref25">Munro and Derwing, 1995</xref>). Critically, L1 effects are highly category-and cue-specific, creating a &#x201C;learnable&#x201D; statistical structure (<xref ref-type="bibr" rid="ref29">Vaughn et al., 2019</xref>; <xref ref-type="bibr" rid="ref34">Xie and Jaeger, 2020</xref>). Indeed, <xref ref-type="bibr" rid="ref16">Eisner et al. (2013)</xref> demonstrated that British English listeners adapted to the word-final devoicing in Dutch-accented English, a finding that has since been extended to other L1-L2 accents (e.g., Mandarin-accented American English, <xref ref-type="bibr" rid="ref38">Xie et al., 2017</xref>). Rapid adaptation has been seen in populations with varying auditory sensitivity and memory capacity (<xref ref-type="bibr" rid="ref7">Bieber and Gordon-Salant, 2017</xref>, <xref ref-type="bibr" rid="ref8">2021</xref>) and can generalize (albeit with limits) across talkers who share an accent (<xref ref-type="bibr" rid="ref2">Baese-Berk et al., 2013</xref>). Exposure benefits in nonnative accent adaptation have thus served as a rich testbed for theories of perceptual learning, adaptation, and its generalization.</p>
<p>Beyond relatively short-term adaptation, real-world speech recognition tends to evolve over repeated episodes of social interaction across talkers and contexts distributed over much longer time spans. Long-term familiarization with an accent over months and years can facilitate the comprehension of, and adaptation to, a novel talker from the same or similar accent background (<xref ref-type="bibr" rid="ref30">Weber et al., 2014</xref>). <xref ref-type="bibr" rid="ref33">Witteman et al. (2013b)</xref> tested native Dutch listeners with limited or extensive prior experience with German-accented Dutch in spoken word recognition. Only those with extended experience with the accent were able to activate the correct lexical entities when hearing heavily accented tokens. <xref ref-type="bibr" rid="ref26">Porretta et al. (2016)</xref> further demonstrated a gradient effect of accent familiarity on the lexical processing of spoken words. From these experimental results, and many personal anecdotes, it is tempting to conclude that repeated exposure accumulates to support adaptive speech perception. However, it is also known that environmental exposure to a previously unfamiliar accent alone does not always lead to a significant change of perception (<xref ref-type="bibr" rid="ref17">Evans and Iverson, 2007</xref>).</p>
<p>Thus, it remains an open question how much exposure could lead to stable, long-lasting perceptual benefits, and existing results are mixed. <xref ref-type="bibr" rid="ref31">Witteman et al. (2015)</xref> showed that adaptation induced by only 3.5&#x2009;min of exposure could be detected as far out as a week later. On the other hand, <xref ref-type="bibr" rid="ref7">Bieber and Gordon-Salant (2017)</xref> found that neither younger nor older adults retained the initial benefit in a delayed test 7&#x2013;10&#x2009;days after exposure (see also <xref ref-type="bibr" rid="ref41">Zheng and Samuel, 2023</xref>). These could be due to differences in methods (e.g., cross-modal priming vs. speech repetition), accent types (e.g., Hebrew-accented Dutch vs. Spanish-accented English), or measures (e.g., adaptation to a single talker vs. generalization across talkers and accents). Regardless, an important gap in the knowledge is whether and if so, how short-term adaptation relates to more long-term changes of perception and/or learning of underlying linguistic representations (for reviews, see <xref ref-type="bibr" rid="ref9">Bieber et al., 2023</xref>; <xref ref-type="bibr" rid="ref35">Xie et al., 2023</xref>).</p>
<p>While there is a growing interest in adaptation across various timescales (<xref ref-type="bibr" rid="ref7">Bieber and Gordon-Salant, 2017</xref>; <xref ref-type="bibr" rid="ref4">Banai et al., 2022</xref>; <xref ref-type="bibr" rid="ref9">Bieber et al., 2023</xref>; <xref ref-type="bibr" rid="ref41">Zheng and Samuel, 2023</xref>), empirical investigation into the medium-term effects&#x2014;spanning days to weeks&#x2014;remains limited. Furthermore, most existing data are from a single, delayed test. They therefore provide little information about the effects of repeated exposure to the same accent, although such interactions are common in real-life social, educational, and workplace settings. Would listeners maintain adaptive changes across these encounters, or would they start over each time? This gap underscores the need for mid-to long-term study designs that more accurately reflect everyday accent exposure and adaptation processes.</p>
<p>Two major challenges remain. The first is subject retention, in particular, to scale up the paradigm to even longer time periods with more frequent tests than the ones considered here. We approached this challenge by using a web-based paradigm that has previously been employed in single-session experiments. Participants are recruited from an online research participant recruitment platform (e.g., Prolific). Making the experiment fully online substantially lowered the effort required by participants (e.g., no need to visit the lab), increasing accessibility and achieving manageable subject attrition (&#x003C;30% over 3&#x2009;weeks). An equivalent in-lab procedure would be more time-consuming for both participants and researchers, which would limit experimental design options, subject eligibility, and retention.</p>
<p>A second challenge&#x2014;one of relevance to research on adaptive speech perception in general&#x2014;is that any test also constitutes a form of exposure, so repeated testing can interfere with researchers&#x2019; ability to accurately measure the effects of exposure. Consider a scenario in which one group is exposed to L2-accented speech, and the other to L1-accented speech. Repeated testing on L2-accented speech tokens inevitably dilutes the difference between the two groups. In other words, the test tokens themselves provide participants with information about the target accent, even in the absence of audio-visual, lexical, or other context that effectively labels the input (<xref ref-type="bibr" rid="ref24">Maye et al., 2002</xref>; <xref ref-type="bibr" rid="ref13">Clayards et al., 2008</xref>). Test stimuli that are often thought of as &#x201C;neutral,&#x201D; such as those sampled uniformly across a continuum, are not free of bias. Listeners can learn the unique statistic in the test stimuli, which gets integrated into and eventually overrides exposure effects. In fact, recent papers provide evidence that prolonged testing reversed adaptive changes that occurred during exposure (<xref ref-type="bibr" rid="ref23">Liu and Jaeger, 2018</xref>; <xref ref-type="bibr" rid="ref9004">Cummings and Theodore, 2023</xref>; <xref ref-type="bibr" rid="ref41">Zheng and Samuel, 2023</xref>).</p>
<p>To address these challenges, we developed a new testing paradigm that balances two competing motivations. On the one hand, we need to keep the number of trials to a minimum to not interfere with exposure. On the other hand, we need to ensure sufficient statistical power to detect the effect of exposure. To achieve this, we created a repeated-exposure-test protocol in which three short (10-item) test blocks alternate with two relatively long exposure blocks within a single session (<xref ref-type="fig" rid="fig1">Figure 1</xref>). We then repeat the five-block design five times. This allows us to keep each test block short, while increasing the total number of test trials and statistical power of the data. In addition, this new protocol enables us to accomplish the overarching goal of tracking the development of adaptation on multiple time scales, from after the first few minutes of exposure to up to 3&#x2009;weeks.</p>
<fig position="float" id="fig1">
<label>Figure 1</label>
<caption>
<p>Each session consisted of three test blocks (10 items each) interspersed with two exposure blocks (90 items each) within a single session. This block design was repeated five times: Day 1, 1&#x2009;day later, 1&#x2009;week later, 2&#x2009;weeks later, and 3&#x2009;weeks later.</p>
</caption>
<graphic xlink:href="fpsyg-15-1383904-g001.tif"/>
</fig>
<p>Our design built on <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref>, who used a single-session exposure-test design to examine L1-English listeners&#x2019; adaptation to a Mandarin-accent word-final /d/-/t/ contrast in English. A syllable-final /d/ vs. /t/ in this accent is often contrasted by burst duration, rather than the cues expected to be most informative in L1-accened English (closure and vowel duration, for more details see 2.2). Many instances of a syllable final /d/ (e.g., &#x201C;kid&#x201D;) sound like a /t/ (e.g., &#x201C;kit&#x201D;) to L1 listeners, and adaptation includes learning to upweight the burst duration over the other cues. That is, rather than examining global improvements independent of accent features, this study zoomed in on how L1 listeners learn new acoustic cue distributions for this specific contrast through exposure to natural Mandarin accents.</p>
<p><xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref> used a between-subjects manipulation in which native listeners of American English answered 180 lexical decision questions in the exposure phase. Participants in the target condition heard 30 (11%) items containing a syllable-final /d/ sound in a lexically-biased context (e.g., &#x201C;lemona<underline>d</underline>e&#x201D;), which was expected to support their adaptation to the accent feature. In the control condition, these items were replaced with words without a syllable-final /d/ sound, and no adaptation was expected. No other stop sounds were present at syllable-final position throughout exposure. During the test, all participants responded to 60 minimal pairs (i.e., 120 items) in a phonetic categorization task (e.g., &#x201C;kid&#x201D; or &#x201C;kit&#x201D;?). They found that exposure to the critical accent feature significantly improved categorization accuracy of /d/, reflecting adaptation to the nonnative accent feature.</p>
<p>As in <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref>, one group of participants in the current experiment were exposed to Mandarin-accented US English, where the critical words contain a syllable-final /d/ sound. Unlike in <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref>, the control participants heard an L1-accented (i.e., native) US English talker producing the same set of words, including the critical words with a syllable-final /d/ sound. Although none of the original exposure items contained stop voicing contrasts, there may be other accent features covarying with the relevant /d/-/t/ contrast (e.g., the realization of voicing in fricatives) which may aid adaptation. This concern is particularly strong for our current multi-session protocol, where a listener receives a large number of exposure trials to a particular talker (180 &#x002A; five sessions&#x2009;=&#x2009;900 trials) over five sessions. We therefore used L1-accented talker in the control condition to remove this possible confound.</p>
<p>During test, both groups were tested on the same set of Mandarin-accented L2 US English tokens which did not occur during exposure. Each test block was brief (10 items from five minimal pairs), and the same set of items were repeated in all test blocks. To counteract the reduced number of test items per block [five (8.3%) out of 60 pairs from <xref ref-type="bibr" rid="ref38">Xie et al., 2017</xref>], we carefully selected stimuli that were predicted to increase statistical power of the results (see details in the Methods section).</p>
<p>We considered two broad classes of results, each of which could shed light on how adaptation develops over days and weeks. If adaptation occurs rapidly but also decays rapidly, benefits originating from the L2-accented (vs. L1-accented) exposure should be found within each session (<xref ref-type="bibr" rid="ref16">Eisner et al., 2013</xref>; <xref ref-type="bibr" rid="ref38">Xie et al., 2017</xref>) but not accumulate across sessions. On the other hand, if immediate and rapid adaptation does lead to enduring changes, the exposure benefits should interact with the number of blocks, e.g., the accuracy difference between the L2-accented vs. L1-accented exposure conditions should increase over the 15 test blocks. Due to the novelty of the paradigm, some of the methodological decisions were made based on related studies. Much of the empirical data needed to formally test hypotheses were not available (e.g., effect sizes and participant attrition rates across multiple sessions). The current methods are thus meant as our initial attempt. In General Discussion, we suggest potential refinements based on the data from this study.</p>
</sec>
<sec sec-type="methods" id="sec2">
<label>2</label>
<title>Methods</title>
<p>All data, analysis scripts, and model summaries are downloadable from OSF (<ext-link xlink:href="https://osf.io/5xfpe/" ext-link-type="uri">osf.io/5xfpe/</ext-link>).</p>
<sec id="sec3">
<label>2.1</label>
<title>Participants</title>
<p>An initial group of 127 monolingual, native speakers of American English, aged 18&#x2013;45, were recruited via Prolific<xref ref-type="fn" rid="fn0001"><sup>1</sup></xref> and completed Session 1 of the experiment via the online testing platform FindingFive.<xref ref-type="fn" rid="fn0002"><sup>2</sup></xref> Because the precise estimates of effect sizes and participant attrition rates were <italic>a priori</italic> unknown, the initial recruitment goal was set based on the published work (<xref ref-type="bibr" rid="ref38">Xie et al., 2017</xref>). The original work included 24 participants in each condition (48 in total); this sample size was equal to or larger than that in other similar work investigating accent adaptation (e.g., <xref ref-type="bibr" rid="ref16">Eisner et al., 2013</xref>; <xref ref-type="bibr" rid="ref32">Witteman et al., 2013a</xref>; <xref ref-type="bibr" rid="ref40">Zheng and Samuel, 2020</xref>). To buffer against subject attrition and increased response variability expected in online testing, we recruited 60&#x2013;65 participants (i.e., approximately 250% increase) in each condition in Session 1.</p>
<p>Due to an administrative error after Session 1, which has subsequently been corrected, 30 participants were unable to continue to the following sessions (see <xref rid="SM1" ref-type="supplementary-material">Supplementary material</xref>). Of the remaining 97 participants, 70 participants (71%) completed all five sessions (<italic>n</italic>&#x2009;=&#x2009;32 in the L2-accented exposure condition; n&#x2009;=&#x2009;38 in the L1-accented exposure condition). Thus, the attrition rate for the last four sessions spanning 3&#x2009;weeks was 27.8%, and comparable between the two exposure conditions: ~28.3% for the L2-accented exposure (15 out of 53) and&#x2009;~&#x2009;27.2% for the L1-accented exposure (12 out of 44).</p>
<p>The 70 participants included in the analysis were recruited from 38 US states, and self-identified as native speakers of US English. Only 6% (four out of 70) of participants reported that they regularly hear Mandarin Chinese spoken by a family member or a close friend; three of them also reported having parents who speak English with a nonnative accent. As we reported in the <xref rid="SM1" ref-type="supplementary-material">Supplementary material</xref>, excluding these four participants did not change the results.</p>
</sec>
<sec id="sec4">
<label>2.2</label>
<title>Stimuli</title>
<p><italic>Exposure stimuli</italic> for the L2-accented exposure group consisted of 90 English words (30 critical and 60 filler items) and 90 phonotactically-legal nonwords. The critical items were all multisyllabic words ending in /d/ (e.g., lemonade, overload). The exposure list for the L1-accented exposure group was identical except that they were produced by a native speaker of American English. Filler words and nonwords did not contain any /d/ or /t/ sounds, and no stop sounds other than /d/ appeared in the word-final position. The exposure items were evenly distributed across the two exposure blocks, each of which thus contained 50% of the exposure stimuli from <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref>. The word-block assignment was counterbalanced across participants and remained constant within participants across the five sessions. Item presentation was randomized within each block.</p>
<p><italic>Test stimuli</italic> consisted of five /d/-/t/-final minimal pairs (e.g., <italic>feed-feet</italic>; 8% of <xref ref-type="bibr" rid="ref38">Xie et al., 2017</xref>), constituting 10 trials per block. This small set of test items was intended to minimize the interference with exposure effects. To increase the chance of detecting adaptation, we selected test tokens that were predicted to be categorized differently after L1-and L2-accented exposure (e.g., a /d/-ending word incorrectly recognized as <italic>_t</italic> after L1-accented exposure but correctly recognized as <italic>_d</italic> after L2-accented exposure). To the extent that past work has taken similar steps, this has typically been focused on the selection of test <italic>talkers</italic> rather than the selection of specific <italic>stimuli</italic>. For example, it is common to select L2 talkers with low-to-medium intelligibility to avoid floor and ceiling effects. However, the effectiveness of an individual stimulus token is known to vary <italic>within</italic> a talker, depending on their exact acoustic-phonetic properties (<xref ref-type="bibr" rid="ref11">Burchill, 2023</xref>; <xref ref-type="bibr" rid="ref35">Xie et al., 2023</xref>).</p>
<p>With this in mind, we first examined the acoustic cues of all 60 pairs of test items used in <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref> across the three cue dimensions: burst, closure, and vowel. We compared them to typical distributions in L1-accented speech, illustrated by blue and yellow ellipses in <xref ref-type="fig" rid="fig2">Figure 2A</xref>. As noted above, the L1 category distributions are primarily separated by closure and vowel duration; in contrast, L2 Mandarin-accented talkers tend to use burst duration, leading to an overlap in the other two dimensions between /d/ and /t/ categories (<xref ref-type="fig" rid="fig2">Figure 2B</xref>) and potential confusion for L1 listeners. While L1 listeners may theoretically resolve this confusion by placing more the perceptual weight on burst duration as a cue for distinguishing /d/ and /t/ sounds, the informativeness of burst duration varies across items due to interactions with the other two cues (<xref ref-type="fig" rid="fig2">Figure 2C</xref>).</p>
<fig position="float" id="fig2">
<label>Figure 2</label>
<caption>
<p><bold>(A)</bold> The selected test items (diamonds) plotted against L1-accented /d/ and /t/ categories in a three-dimensional talker-normalized phonetic space (vowel, closure and burst; for details, see <xref ref-type="bibr" rid="ref27">Tan et al., 2021</xref>). The ellipses show 95% probability density of multivariate Gaussian categories. For details of data used to estimate the distributions, see <xref ref-type="bibr" rid="ref27">Tan et al. (2021)</xref>. <bold>(B)</bold> Same as Panel <bold>(A)</bold>, but seen from a top view, emphasizing the distribution along closure and vowel duration. The selected test items fall into an &#x2018;ambiguous&#x2019; region of the acoustic-phonetic space where L1 /d/ and /t/ overlapped. <bold>(C)</bold> The selected test items (diamonds) plotted against the other 50 test item pairs used in <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref> (circles). While burst duration is generally informative about a given item&#x2019;s category membership (=/d/ or /t/?), the informativeness varies across items. The selected pairs were those predicted by models of distributional learning to yield major improvements in the recognition accuracy after the L2-accented exposure (relative to L1-accented exposure). A contrast between a selected pair (<italic>wet</italic> vs. <italic>wed</italic>) and a pair not selected (<italic>rate</italic> vs. <italic>raid</italic>) is highlighted to illustrate this point.</p>
</caption>
<graphic xlink:href="fpsyg-15-1383904-g002.tif"/>
</fig>
<p>We then used a model to predict how listeners would respond to <italic>each</italic> test token under different exposure conditions, considering all three acoustic cues (vowel, closure, and burst duration). <xref ref-type="bibr" rid="ref27">Tan et al. (2021)</xref> used Bayesian ideal observer models to simulate the outcomes of the two exposure conditions from <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref>. Each model&#x2019;s /d/ category was trained on the respective exposure tokens annotated for the three cues. Since exposure in both groups never included instances of /t/, the /t/ category for both exposure conditions were trained on US English /t/-productions, based on the assumption that L1 listeners in both groups would apply their <italic>a priori</italic> (= L1-based) expectation for the /t/ category. Their results showed that model predictions significantly predicted human categorization responses in each exposure condition at the token level.</p>
<p>We applied the same simulation approach, using MVBeliefUpdatr (<xref ref-type="bibr" rid="ref22">Jaeger and Burchill, 2021</xref>) along with the R code distributed as part of <xref ref-type="bibr" rid="ref35">Xie et al. (2023)</xref>. We ranked all the 60 pairs of test items in terms of the predicted L2-accent exposure advantage (<xref ref-type="fig" rid="fig3">Figure 3</xref>). From the ranked items, we selected five pairs associated with a strong advantage while controlling other factors that would plausibly affect the effectiveness of the test items (e.g., word frequency, vowel types, and ceiling/floor effects on both the /d/ and the /t/ members of a minimal pair). The selected items were thus associated with a significantly higher level of L2-accent exposure advantage, well above the mean of the original 60 pairs.</p>
<fig position="float" id="fig3">
<label>Figure 3</label>
<caption>
<p>/d/ test tokens ranked by the model-predicted accuracy advantage of L2-accented over L1-accented exposure conditions. The tokens selected for the current experiment are highlighted in blue. Vertical dashed lines indicate the average for all the 60 pairs from <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref> (gray) and the five selected pairs (blue).</p>
</caption>
<graphic xlink:href="fpsyg-15-1383904-g003.tif"/>
</fig>
</sec>
<sec id="sec5">
<label>2.3</label>
<title>Procedure</title>
<p>Five experimental sessions were administered on 5&#x2009;days over the course of 3&#x2009;weeks (<xref ref-type="fig" rid="fig1">Figure 1</xref>). Each session consisted of three test blocks (10 trials each) interleaved with two exposure blocks (90 trials each). After a headphone check to adjust volume and confirm the audibility of the audio stimuli, participants began with a test block. Participants were informed that during this block they would hear words ending in /d/ or /t/ and asked to provide two-alternative forced choice (2AFC) responses to the question &#x201C;Did you hear a D or a T?&#x201D; The 10 items from five minimal pairs (e.g., &#x201C;kid&#x201D; or &#x201C;kit&#x201D;) were presented in random order without repetition within each test block.</p>
<p>During exposure, participants completed a lexical decision task (i.e., word or nonword). 180 trials from <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref> were equally divided between two lists presented in the two exposure blocks. Participants heard one token at a time and responded whether it was a real word of English (e.g., &#x201C;lemonade&#x201D;) or a nonword (e.g., &#x201C;salvary&#x201D;). After the five test/exposure blocks, all participants completed a questionnaire about their language background and familiarity with L2 accents in the first session (Day 1). The median participation time was 23&#x2009;min per session.</p>
<p>Using Prolific&#x2019;s &#x201C;custom allow list&#x201D; feature, we invited the participants back to our experiment four more times. We also used Prolific&#x2019;s communication system to send periodic reminders to reduce attrition. During each of Sessions 2&#x2013;5, the experiment was open for 24&#x2009;h starting at 9&#x2009;am Pacific time on a given day, ensuring that the interval between sessions was 1&#x2009;day (sessions 1&#x2013;2) and 1&#x2009;week (after session 2), while the exact interval duration varied across participants. Delayed participation beyond this 24 h window was not permitted. We note that all but three participants across all five sessions completed the experiment between 9&#x2009;am-10&#x2009;pm Pacific time, with the majority completing the experiment in the morning and afternoon before 6&#x2009;pm.</p>
</sec>
</sec>
<sec sec-type="results" id="sec6">
<label>3</label>
<title>Results</title>
<sec id="sec7">
<label>3.1</label>
<title>Exposure</title>
<p><xref ref-type="fig" rid="fig4">Figure 4</xref> shows the overall performance on the lexical decision task during exposure. As expected, participants in the L2-accented exposure group had lower accuracy (mean&#x2009;=&#x2009;0.85; SD&#x2009;=&#x2009;0.03) than the L1-accented exposure group (mean&#x2009;=&#x2009;0.96; SD&#x2009;=&#x2009;0.03). Focusing on the critical /d/-final words, the L2-accented exposure group showed a steady improvement within each session and across sessions (1<sup>st</sup> block; mean&#x2009;=&#x2009;0.78, SD&#x2009;=&#x2009;0.12; last block: mean&#x2009;=&#x2009;0.88, SD&#x2009;=&#x2009;0.11). Meanwhile, performance in the L1-accented exposure group was near ceiling throughout (1<sup>st</sup> block; mean&#x2009;=&#x2009;0.98, SD&#x2009;=&#x2009;0.04; last block: mean&#x2009;=&#x2009;0.98, SD&#x2009;=&#x2009;0.04). The incremental improvement in the L2-accented exposure group suggests that (1) even 15 critical items per exposure block were sufficient to enhance recognition of L2-accented speech, and (2) these enhancements accumulated with increasing exposure.</p>
<fig position="float" id="fig4">
<label>Figure 4</label>
<caption>
<p>Recognition accuracy for the critical /d/-final words in the auditory lexical decision task during exposure blocks spanning 3 weeks. Error bars represent bootstrapped 95% confidence intervals over by-participant means.</p>
</caption>
<graphic xlink:href="fpsyg-15-1383904-g004.tif"/>
</fig>
</sec>
<sec id="sec8">
<label>3.2</label>
<title>Test</title>
<p><xref ref-type="fig" rid="fig5">Figure 5</xref> summarizes participants&#x2019; categorization accuracy on the L2-accented test tokens. As predicted, recognition of /d/-final words (e.g., <italic>feed</italic>) was initially less accurate than recognition of /t/-final words (e.g., <italic>feet</italic>). Also as predicted, recognition accuracy for /d/-final words increased steadily from 0.44 (SD&#x2009;=&#x2009;0.19) on day 1 to 0.62 (SD&#x2009;=&#x2009;0.20) on the final day in week 4 in the L2-accented exposure group, and from 0.41 (SD&#x2009;=&#x2009;0.21) to 0.68 (SD&#x2009;=&#x2009;0.22) in the L1-accented exposure group.</p>
<fig position="float" id="fig5">
<label>Figure 5</label>
<caption>
<p>Performance during the test blocks spanning 3&#x2009;weeks. Error bars represent bootstrapped 95% confidence intervals over by-participant means.</p>
</caption>
<graphic xlink:href="fpsyg-15-1383904-g005.tif"/>
</fig>
<p>We fit a mixed-effect logistic regression (<xref ref-type="bibr" rid="ref21">Jaeger, 2008</xref>) to the test data using the <italic>lme4</italic> package in R (<xref ref-type="bibr" rid="ref5">Bates et al., 2015</xref>). The analysis predicted accuracy (1&#x2009;=&#x2009;correct, 0&#x2009;=&#x2009;incorrect) from the full factorial of exposure condition (effect-coded, -0.5&#x2009;=&#x2009;L1-accented exposure vs. +0.5&#x2009;=&#x2009;L2-accented exposure), category (effect-coded, -0.5 = /t/- vs. +0.5 = /d/-final words), and test block (1&#x2013;15 as a numeric variable, scaled by dividing by two standard deviations, <xref ref-type="bibr" rid="ref19">Gelman, 2008</xref>). Coding test block as a numeric variable allowed us to examine whether incremental, repeated exposure resulted in cumulative improvement in the test performance. We also report in the <xref rid="SM1" ref-type="supplementary-material">Supplementary material</xref> on a separate analysis where we coded test block as an ordered categorical variable. We began with the maximal random effect structure justified by the design and stepwise removed higher-order interactions in the event of convergence failure. The final model included random by-participant intercepts and slopes for category, as well as by-item intercepts and slopes for exposure condition, category, and their interaction.</p>
<p>Both groups&#x2019; overall performance improved significantly across time, as suggested by a significant main effect of test block (<inline-formula>
<mml:math id="M1">
<mml:mover accent="true">
<mml:mi>&#x03B2;</mml:mi>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula> = 0.35, SE&#x2009;=&#x2009;0.05, <italic>z</italic>&#x2009;=&#x2009;6.59, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.0001). The test block-by-category interaction was also significant (<inline-formula>
<mml:math id="M2">
<mml:mover accent="true">
<mml:mi>&#x03B2;</mml:mi>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula> = 0.10, SE&#x2009;=&#x2009;0.11, <italic>z</italic>&#x2009;=&#x2009;9.53, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.0001), indicating that the improvement differed between /d/- and /t/-final words. A follow-up simple effects analysis found recognition accuracy significantly increased for /d/-final words (<inline-formula>
<mml:math id="M3">
<mml:mover accent="true">
<mml:mi>&#x03B2;</mml:mi>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula> = 0.85, SE&#x2009;=&#x2009;0.07, <italic>z</italic>&#x2009;=&#x2009;11.79, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.0001) and significantly <italic>de</italic>creased for /t/-final words (<inline-formula>
<mml:math id="M4">
<mml:mover accent="true">
<mml:mi>&#x03B2;</mml:mi>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula> = &#x2212;0.16, SE&#x2009;=&#x2009;0.08, <italic>z</italic>&#x2009;=&#x2009;&#x2212;2.01, <italic>p</italic>&#x2009;&#x003C;&#x2009;0.05). The overall improvement across test blocks is thus driven by the larger improvements for /d/-final words than the decreased accuracy for /t/-final words (0.85 vs. &#x2013;0.16 log-odds).</p>
<p>To our surprise, and in contrast to an earlier, single-session experiment (<xref ref-type="bibr" rid="ref38">Xie et al., 2017</xref>), no advantage of L2-accented exposure was observed. Neither the main effect of exposure (<inline-formula>
<mml:math id="M5">
<mml:mover accent="true">
<mml:mi>&#x03B2;</mml:mi>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula> = &#x2212;0.01, SE&#x2009;=&#x2009;0.17, <italic>z</italic>&#x2009;=&#x2009;&#x2212;0.07), its two-way interaction with test block (<inline-formula>
<mml:math id="M6">
<mml:mover accent="true">
<mml:mi>&#x03B2;</mml:mi>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula> = &#x2212;0.15, SE =0.10, <italic>z</italic>&#x2009;=&#x2009;&#x2212;1.40), nor its three-way interaction with test block and category (<inline-formula>
<mml:math id="M7">
<mml:mover accent="true">
<mml:mi>&#x03B2;</mml:mi>
<mml:mo stretchy="true">^</mml:mo>
</mml:mover>
</mml:math>
</inline-formula> = &#x2212;0.10, SE&#x2009;=&#x2009;0.21, <italic>z</italic>&#x2009;=&#x2009;&#x2212;0.46) was significant. A post-hoc by-item analysis further confirmed that both groups responded similarly to each of the five minimal pairs across the five sessions (<xref rid="SM1" ref-type="supplementary-material">Supplementary material</xref>).</p>
<p>In summary, we found cumulative improvements in recognition accuracy on the L2-accented exposure and test tokens. This validates the new testing paradigm and demonstrates that it is effective for tracking long-term changes in recognition accuracy. However, contrary to our expectation, this effect did not depend on whether exposure involved L2-or L1-accented speech.</p>
</sec>
</sec>
<sec id="sec9">
<label>4</label>
<title>General discussion</title>
<p>The repeated exposure-test paradigm we explored here holds the potential to bridge two key areas of work on adaptive speech perception: (1) adaptation in the first moments of encountering an (<italic>a priori</italic>) unfamiliar speaker/accent and (2) longitudinal accommodation through repeated environmental exposure spanning weeks, months, and even years. While often assumed, the link between adaptive changes of perception across multiple timescales has rarely been directly tested.</p>
<p>Previous work has shown that exposure-induced changes in speech perception can be detected even up to 1&#x2009;week after exposure (<xref ref-type="bibr" rid="ref15">Eisner and McQueen, 2006</xref>; <xref ref-type="bibr" rid="ref31">Witteman et al., 2015</xref>), albeit sometimes with reduced magnitude (<xref ref-type="bibr" rid="ref41">Zheng and Samuel, 2023</xref>). While these findings speak to the longevity of the adaptive changes in speech perception from even relatively brief exposure, they leave open how <italic>repeated</italic> exposure affects perception. This question is not only of theoretical interest, but also helps to extend scientific knowledge to various ecologically valid scenarios of nonnative speech perception. In real life, listeners often repeatedly encounter a talker and/or an accent over days and weeks. The paradigm we have begun to develop here is meant to simulate this, allowing insights into how accent adaptation develops with repeated exposure (e.g., recurring work calls with international colleagues, listening to a nonnative course instructor).</p>
<p>A recent paper by Bieber and colleagues has begun to address this gap via a unique data set. In their study, eight L1 listeners responded to 750 sentences recorded by 60 different NATO officers, including 44 nonnative talkers from 13 different L1 backgrounds (<xref ref-type="bibr" rid="ref9">Bieber et al., 2023</xref>). Listeners heard 50 sentences per block and identified multiple keywords in each sentence by clicking on written words on a tablet. Critically, they completed a total of 15 blocks in multiple sessions over five to ten days. <xref ref-type="bibr" rid="ref9">Bieber et al. (2023)</xref> found that the greatest degree of benefit from exposure to nonnative accented speech was observed in the first block (~15% increase in accuracy). Performance continued to increase at a slower rate, but with no significant loss of accuracy between sessions (with intervals of 1&#x2013;4&#x2009;days between sessions). This suggests that adaptation to nonnative accents can be maintained, and it accumulates over time. However, the small sample size (eight listeners) and the relatively short duration (up to 10&#x2009;days) leave open the question of how adaptation may proceed when the subject pool is more heterogeneous, and exposure is more widely spaced.</p>
<p>Encouragingly, our findings demonstrated the feasibility of longitudinal studies with larger participant groups tested on an online platform. Over 3&#x2009;weeks, we found that sustainable subject retention is possible: After the initial technical issues we encountered (avoidable in future applications of the paradigm), subject retention was high. The majority of participants who committed to the first two sessions completed all five sessions. Moreover, while the behavioral improvements continued over time, the recognition accuracy after the fifth session was still far from ceiling. This pattern of results suggests the potential for using this paradigm to explore longer-term adaptive changes in perception, possibly over months. Importantly, the current study is one of the first to examine long-term changes in the recognition of specific phonetic categories, beyond general improvements in accented speech recognition (e.g., <xref ref-type="bibr" rid="ref9">Bieber et al., 2023</xref>). This opens avenues for research on listeners&#x2019; adaptation to underlying phonetic category representations as a driver of longitudinal perceptual change.</p>
<p>However, the failure to replicate the advantage of L2-accented exposure found in previous single-session experiments also points to a challenge for similar future studies. We highlight three possible reasons for this unexpected result.</p>
<p>First, it is possible that it was a Type II error. One question is thus whether the present experiment was under-powered compared to <xref ref-type="bibr" rid="ref38">Xie et al. (2017)</xref>. On the one hand, each of our tests employed substantially fewer test tokens compared to the original single-session study (five pairs instead of 60). On the other hand, the simulation-based stimuli selection (Section 2.2.) countered this loss of power, and the present experiment employed substantially <italic>more</italic> test blocks (15 instead of one). A look at participant numbers is similarly uninformative: while the unexpected loss of participants after Session 1 reduced the participants available for analysis, it still left us with more participants than the original study (70 web-based vs. 48 lab-based). None of these considerations thus point to a clear power disadvantage of the present study. Still, to empirically address this question, we conducted a power simulation (for details, see <xref rid="SM1" ref-type="supplementary-material">Supplementary material</xref>). These simulations indeed estimated over 95% power to detect an advantage of L2-accented exposure with the effect sizes from the original study.</p>
<p>A second possible explanation for the failure to detect an advantage of L2-accented exposure concerns differences in the design. Specifically, the &#x201C;control&#x201D; condition in a related work (<xref ref-type="bibr" rid="ref16">Eisner et al., 2013</xref>; <xref ref-type="bibr" rid="ref38">Xie et al., 2017</xref>) typically employed L2-accented speech without a critical /d/-final word. In contrast, the current control condition used L1-accented speech with /d/-final word present. This design may have helped the L1-accented exposure group directly contrast the L1 and L2 accents and isolate the critical phonetic differences (e.g., <xref ref-type="bibr" rid="ref14">Cooper and Bradlow, 2016</xref>). Under this explanation, participants in the two exposure groups both improved their L2-accent recognition but for different reasons. A follow-up experiment to address this possibility is currently underway (<xref ref-type="bibr" rid="ref9002">Kurumada and Xie, in prep</xref>).</p>
<p>Another, mutually compatible, possibility is that the test tokens alone may have been sufficient to support adaptation. That is, participants in the L1-accented exposure group adapted to the accent-specific features through the minimal /d/-/t/ pairs heard during the test. Even though these items were unlabeled, the underlying acoustic features relevant to /d/ vs. /t/ recognition show a natural bimodal distribution along the burst dimension (<xref ref-type="fig" rid="fig2">Figure 2</xref>). Learning from bimodal unlabeled input has been found in previous work, though those studies involved many more tokens (e.g., <xref ref-type="bibr" rid="ref24">Maye et al., 2002</xref>; <xref ref-type="bibr" rid="ref13">Clayards et al., 2008</xref>; <xref ref-type="bibr" rid="ref9001">Kleinschmidt and Jaeger, 2015</xref>; <xref ref-type="bibr" rid="ref9003">Theodore and Monto, 2019</xref>). Additionally, the repeated encounter to the same set of minimal pairs could have endorsed some response strategies. It is therefore important to avoid minimal pairs or select a greater variety of test tokens (while keeping each test block short) to avoid repetition or anchoring of test tokens within and across test blocks.</p>
<p>In summary, we presented a new repeated exposure-test paradigm to investigate the adaptive speech perception over three weeks. The current results and the information derived from the current work lay the empirical ground for future related lines of inquiry. For interested readers, we have provided in the <xref rid="SM1" ref-type="supplementary-material">Supplementary material</xref> insights we gained about the administration of a longitudinal study using an online testing platform.</p>
</sec>
<sec sec-type="data-availability" id="sec10">
<title>Data availability statement</title>
<p>The datasets presented in this study can be found in online repositories. The names of the repository/repositories and accession number(s) can be found in the article/<xref rid="SM1" ref-type="supplementary-material">Supplementary material</xref>.</p>
</sec>
<sec sec-type="ethics-statement" id="sec11">
<title>Ethics statement</title>
<p>The studies involving humans were approved by Institutional Review Board University of California Irvine. The studies were conducted in accordance with the local legislation and institutional requirements. The participants provided their written informed consent to participate in this study.</p>
</sec>
<sec sec-type="author-contributions" id="sec12">
<title>Author contributions</title>
<p>XX: Conceptualization, Data curation, Formal analysis, Funding acquisition, Investigation, Methodology, Project administration, Resources, Software, Supervision, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. CK: Conceptualization, Data curation, Formal analysis, Funding acquisition, Investigation, Methodology, Project administration, Resources, Software, Supervision, Validation, Visualization, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing.</p>
</sec>
</body>
<back>
<sec sec-type="funding-information" id="sec13">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research, authorship, and/or publication of this article. This work was supported by NIH National Institutes of Health (NIH)-National Institute of Child Health and Human Development (NICHD) Grant No. R01HD111936 to XX and CK.</p>
</sec>
<ack>
<p>We thank Haleh Farahbod for her assistance in implementing and administering the experiment, and T. Florian Jaeger, Arthur G. Samuel, Isabella Franchesca Dela Cr Eclevia, Amelia Ostrow, and participants of the 184th ASA meeting and the 20th ICPhS meeting for their helpful comments and suggestions.</p>
</ack>
<sec sec-type="COI-statement" id="sec14">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="sec15">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<sec sec-type="supplementary-material" id="sec16">
<title>Supplementary material</title>
<p>The Supplementary material for this article can be found online at: <ext-link xlink:href="https://www.frontiersin.org/articles/10.3389/fpsyg.2024.1383904/full#supplementary-material" ext-link-type="uri">https://www.frontiersin.org/articles/10.3389/fpsyg.2024.1383904/full#supplementary-material</ext-link></p>
<supplementary-material xlink:href="Data_Sheet_1.pdf" id="SM1" mimetype="application/pdf" xmlns:xlink="http://www.w3.org/1999/xlink"/>
</sec>
<fn-group>
<fn id="fn0001">
<p><sup>1</sup><ext-link xlink:href="https://www.prolific.co/" ext-link-type="uri">https://www.prolific.co/</ext-link>
</p>
</fn>
<fn id="fn0002">
<p><sup>2</sup><ext-link xlink:href="https://www.findingfive.com/" ext-link-type="uri">https://www.findingfive.com/</ext-link>
</p>
</fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="ref1">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Adank</surname> <given-names>P.</given-names></name> <name><surname>Janse</surname> <given-names>E.</given-names></name></person-group> (<year>2010</year>). <article-title>Comprehension of a novel accent by young and older listeners</article-title>. <source>Psychol. Aging</source> <volume>25</volume>, <fpage>736</fpage>&#x2013;<lpage>740</lpage>. doi: <pub-id pub-id-type="doi">10.1037/a0020054</pub-id>, PMID: <pub-id pub-id-type="pmid">20853978</pub-id></citation>
</ref>
<ref id="ref2">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Baese-Berk</surname> <given-names>M. M.</given-names></name> <name><surname>Bradlow</surname> <given-names>A. R.</given-names></name> <name><surname>Wright</surname> <given-names>B. A.</given-names></name></person-group> (<year>2013</year>). <article-title>Accent-independent adaptation to foreign accented speech</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>133</volume>, <fpage>EL174</fpage>&#x2013;<lpage>EL180</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.4789864</pub-id></citation>
</ref>
<ref id="ref3">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Baese-Berk</surname> <given-names>M. M.</given-names></name> <name><surname>McLaughlin</surname> <given-names>D. J.</given-names></name> <name><surname>McGowan</surname> <given-names>K. B.</given-names></name></person-group> (<year>2020</year>). <article-title>Perception of nonnative speech</article-title>. <source>Lang. Linguist. Compass.</source> <volume>14</volume>, <fpage>1</fpage>&#x2013;<lpage>20</lpage>. doi: <pub-id pub-id-type="doi">10.1111/lnc3.12375</pub-id></citation>
</ref>
<ref id="ref4">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Banai</surname> <given-names>K.</given-names></name> <name><surname>Karawani</surname> <given-names>H.</given-names></name> <name><surname>Lavie</surname> <given-names>L.</given-names></name> <name><surname>Lavner</surname> <given-names>Y.</given-names></name></person-group> (<year>2022</year>). <article-title>Rapid but specific perceptual learning partially explains individual differences in the recognition of challenging speech</article-title>. <source>Sci. Rep.</source> <volume>12</volume>:<fpage>10011</fpage>. doi: <pub-id pub-id-type="doi">10.1038/s41598-022-14189-8</pub-id>, PMID: <pub-id pub-id-type="pmid">35705680</pub-id></citation>
</ref>
<ref id="ref5">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bates</surname> <given-names>D.</given-names></name> <name><surname>M&#x00E4;chler</surname> <given-names>M.</given-names></name> <name><surname>Bolker</surname> <given-names>B.</given-names></name> <name><surname>Walker</surname> <given-names>S.</given-names></name></person-group> (<year>2015</year>). <article-title>Fitting linear mixed-effects models using lme4</article-title>. <source>J. Stat. Softw.</source> <volume>67</volume>, <fpage>1</fpage>&#x2013;<lpage>48</lpage>. doi: <pub-id pub-id-type="doi">10.18637/jss.v067.i01</pub-id></citation>
</ref>
<ref id="ref6">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Bent</surname> <given-names>T.</given-names></name> <name><surname>Baese-Berk</surname> <given-names>M.</given-names></name></person-group> (<year>2021</year>). &#x201C;<article-title>Perceptual learning of accented speech</article-title>&#x201D; in <source>The Handbook of Speech Perception</source>. eds. <person-group person-group-type="editor"><name><surname>Pisoni</surname> <given-names>D. B.</given-names></name> <name><surname>Remez</surname> <given-names>R. E.</given-names></name></person-group> (<publisher-loc>Hoboken, NJ</publisher-loc>: <publisher-name>Wiley</publisher-name>).</citation>
</ref>
<ref id="ref7">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bieber</surname> <given-names>R. E.</given-names></name> <name><surname>Gordon-Salant</surname> <given-names>S.</given-names></name></person-group> (<year>2017</year>). <article-title>Adaptation to novel foreign-accented speech and retention of benefit following training: influence of aging and hearing loss</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>141</volume>, <fpage>2800</fpage>&#x2013;<lpage>2811</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.4980063</pub-id>, PMID: <pub-id pub-id-type="pmid">28464671</pub-id></citation>
</ref>
<ref id="ref8">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bieber</surname> <given-names>R. E.</given-names></name> <name><surname>Gordon-Salant</surname> <given-names>S.</given-names></name></person-group> (<year>2021</year>). <article-title>Improving older adults&#x2019; understanding of challenging speech: auditory training, rapid adaptation and perceptual learning</article-title>. <source>Hear. Res.</source> <volume>402</volume>:<fpage>108054</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.heares.2020.108054</pub-id>, PMID: <pub-id pub-id-type="pmid">32826108</pub-id></citation>
</ref>
<ref id="ref9">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bieber</surname> <given-names>R. E.</given-names></name> <name><surname>Makashay</surname> <given-names>M. J.</given-names></name> <name><surname>Simpson</surname> <given-names>B.</given-names></name> <name><surname>Sheffield</surname> <given-names>B. M.</given-names></name> <name><surname>Brungart</surname> <given-names>D. S.</given-names></name></person-group> (<year>2023</year>). <article-title>Short-term retention of learning after rapid adaptation to native and nonnative speech</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>153</volume>:<fpage>3362</fpage>. doi: <pub-id pub-id-type="doi">10.1121/10.0019749</pub-id>, PMID: <pub-id pub-id-type="pmid">37338291</pub-id></citation>
</ref>
<ref id="ref10">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bradlow</surname> <given-names>A. R.</given-names></name> <name><surname>Bent</surname> <given-names>T.</given-names></name></person-group> (<year>2008</year>). <article-title>Perceptual adaptation to nonnative speech</article-title>. <source>Cognition</source> <volume>106</volume>, <fpage>707</fpage>&#x2013;<lpage>729</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.cognition.2007.04.005</pub-id>, PMID: <pub-id pub-id-type="pmid">17532315</pub-id></citation>
</ref>
<ref id="ref11">
<citation citation-type="book"><person-group person-group-type="author">
<name><surname>Burchill</surname> <given-names>Z.</given-names></name>
</person-group> (<year>2023</year>). <source>Understanding the nature of maintained information in speech processing and the reliability of standard reading time analyses</source>. <publisher-loc>Rochester, NY</publisher-loc>: <publisher-name>University of Rochester</publisher-name>.</citation>
</ref>
<ref id="ref12">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Clarke</surname> <given-names>C. M.</given-names></name> <name><surname>Garrett</surname> <given-names>M. F.</given-names></name></person-group> (<year>2004</year>). <article-title>Rapid adaptation to foreign-accented English</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>116</volume>, <fpage>3647</fpage>&#x2013;<lpage>3658</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.1815131</pub-id>, PMID: <pub-id pub-id-type="pmid">15658715</pub-id></citation>
</ref>
<ref id="ref13">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Clayards</surname> <given-names>M.</given-names></name> <name><surname>Tanenhaus</surname> <given-names>M. K.</given-names></name> <name><surname>Aslin</surname> <given-names>R.</given-names></name> <name><surname>Jacobs</surname> <given-names>R. A.</given-names></name></person-group> (<year>2008</year>). <article-title>Perception of speech reflects optimal use of probablistic speech cues</article-title>. <source>Cognition</source> <volume>108</volume>, <fpage>804</fpage>&#x2013;<lpage>809</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.cognition.2008.04.004</pub-id>, PMID: <pub-id pub-id-type="pmid">18582855</pub-id></citation>
</ref>
<ref id="ref14">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cooper</surname> <given-names>A.</given-names></name> <name><surname>Bradlow</surname> <given-names>A. R.</given-names></name></person-group> (<year>2016</year>). <article-title>Linguistically guided adaptation to foreign-accented speech</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>140</volume>, <fpage>EL378</fpage>&#x2013;<lpage>EL384</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.4966585</pub-id>, PMID: <pub-id pub-id-type="pmid">27908062</pub-id></citation>
</ref>
<ref id="ref9004">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Cummings</surname> <given-names>S. N.</given-names></name> <name><surname>Theodore</surname> <given-names>R. M.</given-names></name></person-group> (<year>2023</year>). <article-title>Hearing is believing: Lexically guided perceptual learning is graded to reflect the quantity of evidence in speech input</article-title>. <source>Cognition</source> <volume>235</volume>:<fpage>105404</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.cognition.2023.105404</pub-id></citation>
</ref>
<ref id="ref15">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Eisner</surname> <given-names>F.</given-names></name> <name><surname>McQueen</surname> <given-names>J. M.</given-names></name></person-group> (<year>2006</year>). <article-title>Perceptual learning in speech: stability over time</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>119</volume>, <fpage>1950</fpage>&#x2013;<lpage>1953</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.2178721</pub-id>, PMID: <pub-id pub-id-type="pmid">16642808</pub-id></citation>
</ref>
<ref id="ref16">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Eisner</surname> <given-names>F.</given-names></name> <name><surname>Melinger</surname> <given-names>A.</given-names></name> <name><surname>Weber</surname> <given-names>A.</given-names></name></person-group> (<year>2013</year>). <article-title>Constraints on the transfer of perceptual learning in accented speech</article-title>. <source>Front. Psychol.</source> <volume>4</volume>:<fpage>148</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2013.00148</pub-id>, PMID: <pub-id pub-id-type="pmid">23554598</pub-id></citation>
</ref>
<ref id="ref17">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Evans</surname> <given-names>B. G.</given-names></name> <name><surname>Iverson</surname> <given-names>P.</given-names></name></person-group> (<year>2007</year>). <article-title>Plasticity in vowel perception and production: a study of accent change in young adults</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>121</volume>, <fpage>3814</fpage>&#x2013;<lpage>3826</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.2722209</pub-id>, PMID: <pub-id pub-id-type="pmid">17552729</pub-id></citation>
</ref>
<ref id="ref18">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Flege</surname> <given-names>J. E.</given-names></name> <name><surname>Munro</surname> <given-names>M.</given-names></name> <name><surname>Skelton</surname> <given-names>L.</given-names></name></person-group> (<year>1992</year>). <article-title>Production of the word-final English /t/ - /d/ contrast speakers of English, Mandarin, and Spanish</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>92</volume>, <fpage>128</fpage>&#x2013;<lpage>143</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.404278</pub-id></citation>
</ref>
<ref id="ref19">
<citation citation-type="journal"><person-group person-group-type="author">
<name><surname>Gelman</surname> <given-names>A.</given-names></name>
</person-group> (<year>2008</year>). <article-title>Scaling regression inputs by dividing by two standard deviations</article-title>. <source>Stat. Med.</source> <volume>27</volume>, <fpage>2865</fpage>&#x2013;<lpage>2873</lpage>. doi: <pub-id pub-id-type="doi">10.1002/sim.3107</pub-id></citation>
</ref>
<ref id="ref20">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hanul&#x00ED;kov&#x00E1;</surname> <given-names>A.</given-names></name> <name><surname>Weber</surname> <given-names>A.</given-names></name></person-group> (<year>2012</year>). <article-title>Sink positive: linguistic experience with &#x201C;th&#x201D; substitutions influences nonnative word recognition</article-title>. <source>Atten. Percept. Psychophys.</source> <volume>74</volume>, <fpage>613</fpage>&#x2013;<lpage>629</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13414-011-0259-7</pub-id>, PMID: <pub-id pub-id-type="pmid">22207311</pub-id></citation>
</ref>
<ref id="ref21">
<citation citation-type="journal"><person-group person-group-type="author">
<name><surname>Jaeger</surname> <given-names>T. F.</given-names></name>
</person-group> (<year>2008</year>). <article-title>Categorical data analysis: away from ANOVAs (transformation or not) and towards logit mixed models</article-title>. <source>J. Mem. Lang.</source> <volume>59</volume>, <fpage>434</fpage>&#x2013;<lpage>446</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.jml.2007.11.007</pub-id>, PMID: <pub-id pub-id-type="pmid">19884961</pub-id></citation>
</ref>
<ref id="ref22">
<citation citation-type="other"><person-group person-group-type="author"><name><surname>Jaeger</surname> <given-names>T. F.</given-names></name> <name><surname>Burchill</surname> <given-names>Z.</given-names></name></person-group> (<year>2021</year>). MVBeliefUpdatr. Available at: <ext-link xlink:href="https://github.com/hlplab/MVBeliefUpdatr" ext-link-type="uri">https://github.com/hlplab/MVBeliefUpdatr</ext-link></citation>
</ref>
<ref id="ref9001">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kleinschmidt</surname> <given-names>D. F.</given-names></name> <name><surname>Jaeger</surname> <given-names>T. F.</given-names></name></person-group> (<year>2015</year>). <article-title>Robust speech perception: Recognize the familiar, generalize to the similar, and adapt to the novel</article-title>. <source>Psychol. Rev.</source> <volume>122</volume>, <fpage>148</fpage>&#x2013;<lpage>203</lpage>. doi: <pub-id pub-id-type="doi">10.1037/a0038695</pub-id></citation>
</ref>
<ref id="ref9002">
<citation citation-type="other"><person-group person-group-type="author"><name><surname>Kurumada</surname> <given-names>C.</given-names></name> <name><surname>Xie</surname> <given-names>X</given-names></name></person-group>. (<year>in prep</year>). <article-title>Roles of unlabeled minimal pairs in perceptual adaptation experiments</article-title>.</citation>
</ref>
<ref id="ref23">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>L.</given-names></name> <name><surname>Jaeger</surname> <given-names>T. F.</given-names></name></person-group> (<year>2018</year>). <article-title>Inferring causes during speech perception</article-title>. <source>Cognition</source> <volume>174</volume>, <fpage>55</fpage>&#x2013;<lpage>70</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.cognition.2018.01.003</pub-id>, PMID: <pub-id pub-id-type="pmid">29425987</pub-id></citation>
</ref>
<ref id="ref24">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Maye</surname> <given-names>J.</given-names></name> <name><surname>Werker</surname> <given-names>J. F.</given-names></name> <name><surname>Gerken</surname> <given-names>L.</given-names></name></person-group> (<year>2002</year>). <article-title>Infant sensitivity to distributional information can affect phonetic discrimination</article-title>. <source>Cognition</source> <volume>82</volume>, <fpage>B101</fpage>&#x2013;<lpage>B111</lpage>. doi: <pub-id pub-id-type="doi">10.1016/S0010-0277(01)00157-3</pub-id>, PMID: <pub-id pub-id-type="pmid">11747867</pub-id></citation>
</ref>
<ref id="ref25">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Munro</surname> <given-names>M. J.</given-names></name> <name><surname>Derwing</surname> <given-names>T. M.</given-names></name></person-group> (<year>1995</year>). <article-title>Foreign accent, comprehensibility, and intelligibility in the speech of second language learners</article-title>. <source>Lang. Learn.</source> <volume>45</volume>, <fpage>73</fpage>&#x2013;<lpage>97</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1467-1770.1995.tb00963.x</pub-id></citation>
</ref>
<ref id="ref26">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Porretta</surname> <given-names>V.</given-names></name> <name><surname>Tucker</surname> <given-names>B.</given-names></name> <name><surname>J&#x00E4;rvikivi</surname> <given-names>J.</given-names></name></person-group> (<year>2016</year>). <article-title>The influence of gradient foreign accentedness and listener experience on word recognition</article-title>. <source>J. Phon.</source> <volume>58</volume>, <fpage>1</fpage>&#x2013;<lpage>21</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.wocn.2016.05.006</pub-id></citation>
</ref>
<ref id="ref27">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tan</surname> <given-names>M.</given-names></name> <name><surname>Xie</surname> <given-names>X.</given-names></name> <name><surname>Jaeger</surname> <given-names>T. F.</given-names></name></person-group> (<year>2021</year>). <article-title>Using rational models to interpret the results of experiments on accent adaptation</article-title>. <source>Front. Psychol.</source> <volume>12</volume>:<fpage>676271</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2021.676271</pub-id>, PMID: <pub-id pub-id-type="pmid">34803790</pub-id></citation>
</ref>
<ref id="ref9003">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Theodore</surname> <given-names>R. M.</given-names></name> <name><surname>Monto</surname> <given-names>N. R.</given-names></name></person-group> (<year>2019</year>). <article-title>Distributional learning for speech reflects cumulative exposure to a talker&#x2019;s phonetic distributions</article-title>. <source>Psychon. Bull. Rev.</source> <volume>26</volume>, <fpage>985</fpage>&#x2013;<lpage>992</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13423-018-1551-5</pub-id></citation>
</ref>
<ref id="ref28">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tzeng</surname> <given-names>C. Y.</given-names></name> <name><surname>Alexander</surname> <given-names>J. E. D.</given-names></name> <name><surname>Sidaras</surname> <given-names>S. K.</given-names></name> <name><surname>Nygaard</surname> <given-names>L. C.</given-names></name></person-group> (<year>2016</year>). <article-title>The role of training structure in perceptual learning of accented speech</article-title>. <source>J. Exp. Psychol. Hum. Percept. Perform.</source> <volume>42</volume>, <fpage>1793</fpage>&#x2013;<lpage>1805</lpage>. doi: <pub-id pub-id-type="doi">10.1037/xhp0000260</pub-id>, PMID: <pub-id pub-id-type="pmid">27399829</pub-id></citation>
</ref>
<ref id="ref29">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Vaughn</surname> <given-names>C.</given-names></name> <name><surname>Baese-Berk</surname> <given-names>M.</given-names></name> <name><surname>Idemaru</surname> <given-names>K.</given-names></name></person-group> (<year>2019</year>). <article-title>Re-examining phonetic variability in native and nonnative speech</article-title>. <source>Phonetica</source> <volume>76</volume>, <fpage>327</fpage>&#x2013;<lpage>358</lpage>. doi: <pub-id pub-id-type="doi">10.1159/000487269</pub-id>, PMID: <pub-id pub-id-type="pmid">30086539</pub-id></citation>
</ref>
<ref id="ref30">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Weber</surname> <given-names>A.</given-names></name> <name><surname>Betta</surname> <given-names>A. M.</given-names></name> <name><surname>McQueen</surname> <given-names>J. M.</given-names></name></person-group> (<year>2014</year>). <article-title>Treack or trit: adaptation to genuine and arbitrary foreign accents by monolingual and bilingual listeners</article-title>. <source>J. Phon.</source> <volume>46</volume>, <fpage>34</fpage>&#x2013;<lpage>51</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.wocn.2014.05.002</pub-id></citation>
</ref>
<ref id="ref31">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Witteman</surname> <given-names>M. J.</given-names></name> <name><surname>Bardhan</surname> <given-names>N. P.</given-names></name> <name><surname>Weber</surname> <given-names>A.</given-names></name> <name><surname>McQueen</surname> <given-names>J. M.</given-names></name></person-group> (<year>2015</year>). <article-title>Automaticity and stability of adaptation to a foreign-accented speaker</article-title>. <source>Lang. Speech</source> <volume>58</volume>, <fpage>168</fpage>&#x2013;<lpage>189</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0023830914528102</pub-id>, PMID: <pub-id pub-id-type="pmid">26677641</pub-id></citation>
</ref>
<ref id="ref32">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Witteman</surname> <given-names>M. J.</given-names></name> <name><surname>Weber</surname> <given-names>A.</given-names></name> <name><surname>McQueen</surname> <given-names>J. M.</given-names></name></person-group> (<year>2013a</year>). <article-title>Foreign accent strength and listener familiarity with an accent codetermine speed of perceptual adaptation</article-title>. <source>Atten. Percept. Psychophys.</source> <volume>75</volume>, <fpage>537</fpage>&#x2013;<lpage>556</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13414-012-0404-y</pub-id>, PMID: <pub-id pub-id-type="pmid">23456266</pub-id></citation>
</ref>
<ref id="ref33">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Witteman</surname> <given-names>M. J.</given-names></name> <name><surname>Weber</surname> <given-names>A.</given-names></name> <name><surname>McQueen</surname> <given-names>J. M.</given-names></name></person-group> (<year>2013b</year>). <article-title>Tolerance for inconsistency in foreign-accented speech</article-title>. <source>Psychon. Bull. Rev.</source> <volume>21</volume>, <fpage>512</fpage>&#x2013;<lpage>519</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13423-013-0519-8</pub-id></citation>
</ref>
<ref id="ref34">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xie</surname> <given-names>X.</given-names></name> <name><surname>Jaeger</surname> <given-names>T. F.</given-names></name></person-group> (<year>2020</year>). <article-title>Comparing nonnative and native speech: are L2 productions more variable?</article-title> <source>J. Acoust. Soc. Am.</source> <volume>147</volume>, <fpage>3322</fpage>&#x2013;<lpage>3347</lpage>. doi: <pub-id pub-id-type="doi">10.1121/10.0001141</pub-id>, PMID: <pub-id pub-id-type="pmid">32486781</pub-id></citation>
</ref>
<ref id="ref35">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xie</surname> <given-names>X.</given-names></name> <name><surname>Jaeger</surname> <given-names>T. F.</given-names></name> <name><surname>Kurumada</surname> <given-names>C.</given-names></name></person-group> (<year>2023</year>). <article-title>What we do (not) know about the mechanisms underlying adaptive speech perception: a computational review</article-title>. <source>Cortex</source> <volume>166</volume>, <fpage>377</fpage>&#x2013;<lpage>424</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.cortex.2023.05.003</pub-id>, PMID: <pub-id pub-id-type="pmid">37506665</pub-id></citation>
</ref>
<ref id="ref36">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xie</surname> <given-names>X.</given-names></name> <name><surname>Liu</surname> <given-names>L.</given-names></name> <name><surname>Jaeger</surname> <given-names>T. F.</given-names></name></person-group> (<year>2021</year>). <article-title>Cross-talker generalization in the perception of nonnative speech: a large-scale replication</article-title>. <source>J. Exp. Psychol. Gen.</source> <volume>150</volume>, <fpage>e22</fpage>&#x2013;<lpage>e56</lpage>. doi: <pub-id pub-id-type="doi">10.1037/xge0001039</pub-id></citation>
</ref>
<ref id="ref38">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xie</surname> <given-names>X.</given-names></name> <name><surname>Theodore</surname> <given-names>R. M.</given-names></name> <name><surname>Myers</surname> <given-names>E. B.</given-names></name></person-group> (<year>2017</year>). <article-title>More than a boundary shift: perceptual adaptation to foreign-accented speech reshapes the internal structure of phonetic categories</article-title>. <source>J. Exp. Psychol. Hum. Percept. Perform.</source> <volume>43</volume>, <fpage>206</fpage>&#x2013;<lpage>217</lpage>. doi: <pub-id pub-id-type="doi">10.1037/xhp0000285</pub-id>, PMID: <pub-id pub-id-type="pmid">27819457</pub-id></citation>
</ref>
<ref id="ref39">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xie</surname> <given-names>X.</given-names></name> <name><surname>Weatherholtz</surname> <given-names>K.</given-names></name> <name><surname>Bainton</surname> <given-names>L.</given-names></name> <name><surname>Rowe</surname> <given-names>E.</given-names></name> <name><surname>Burchill</surname> <given-names>Z.</given-names></name> <name><surname>Liu</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Rapid adaptation to foreign-accented speech and its transfer to an unfamiliar talker</article-title>. <source>J. Acoust. Soc. Am.</source> <volume>143</volume>, <fpage>2013</fpage>&#x2013;<lpage>2031</lpage>. doi: <pub-id pub-id-type="doi">10.1121/1.5027410</pub-id>, PMID: <pub-id pub-id-type="pmid">29716296</pub-id></citation>
</ref>
<ref id="ref40">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zheng</surname> <given-names>Y.</given-names></name> <name><surname>Samuel</surname> <given-names>A. G.</given-names></name></person-group> (<year>2020</year>). <article-title>The relationship between phonemic category boundary changes and perceptual adjustments to natural accents</article-title>. <source>J. Exp. Psychol. Learn. Mem. Cogn.</source> <volume>46</volume>, <fpage>1270</fpage>&#x2013;<lpage>1292</lpage>. doi: <pub-id pub-id-type="doi">10.1037/xlm0000788</pub-id>, PMID: <pub-id pub-id-type="pmid">31633368</pub-id></citation>
</ref>
<ref id="ref41">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zheng</surname> <given-names>Y.</given-names></name> <name><surname>Samuel</surname> <given-names>A. G.</given-names></name></person-group> (<year>2023</year>). <article-title>Flexibility and stability of speech sounds: the time course of lexically-driven recalibration</article-title>. <source>J. Phon.</source> <volume>97</volume>:<fpage>101222</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.wocn.2023.101222</pub-id></citation>
</ref>
</ref-list>
</back>
</article>