<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Hum. Neurosci.</journal-id>
<journal-title>Frontiers in Human Neuroscience</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Hum. Neurosci.</abbrev-journal-title>
<issn pub-type="epub">1662-5161</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fnhum.2024.1410242</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Human Neuroscience</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Embodied cognition and L2 sentence comprehension: an eye-tracking study of motor representations</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name><surname>Shiang</surname> <given-names>Ruei-Fang</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x0002A;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/2702181/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Chern</surname> <given-names>Chiou-Lan</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" equal-contrib="yes">
<name><surname>Chen</surname> <given-names>Hsueh-Chih</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="author-notes" rid="fn001"><sup>&#x02020;</sup></xref>
<uri xlink:href="http://loop.frontiersin.org/people/279756/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Department of English, National Taiwan Normal University</institution>, <addr-line>Taipei City</addr-line>, <country>Taiwan</country></aff>
<aff id="aff2"><sup>2</sup><institution>Department of Educational Psychology and Counseling, Chinese Language and Technology Center, Social Emotional Education and Development Center, Institute for Research Excellence in Learning Sciences, National Taiwan Normal University</institution>, <addr-line>Taipei City</addr-line>, <country>Taiwan</country></aff>
<author-notes>
<fn fn-type="edited-by"><p>Edited by: Norbert Vanek, The University of Auckland, New Zealand</p></fn>
<fn fn-type="edited-by"><p>Reviewed by: Yufen Wei, Bangor University, United Kingdom</p>
<p>Yuyan Xue, University of Cambridge, United Kingdom</p></fn>
<corresp id="c001">&#x0002A;Correspondence: Ruei-Fang Shiang <email>joyceshiang&#x00040;ntnu.edu.tw</email></corresp>
<fn fn-type="equal" id="fn001"><p>&#x02020;These authors have contributed equally to this work and share first authorship</p></fn></author-notes>
<pub-date pub-type="epub">
<day>17</day>
<month>09</month>
<year>2024</year>
</pub-date>
<pub-date pub-type="collection">
<year>2024</year>
</pub-date>
<volume>18</volume>
<elocation-id>1410242</elocation-id>
<history>
<date date-type="received">
<day>31</day>
<month>03</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>15</day>
<month>08</month>
<year>2024</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x000A9; 2024 Shiang, Chern and Chen.</copyright-statement>
<copyright-year>2024</copyright-year>
<copyright-holder>Shiang, Chern and Chen</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p></license>
</permissions>
<abstract>
<sec>
<title>Introduction</title>
<p>Evidence from neuroscience and behavioral research has indicated that language meaning is grounded in our motor&#x02013;perceptual experiences of the world. However, the question of whether motor embodiment occurs at the sentence level in L2 (second language) comprehension has been raised. Furthermore, existing studies on motor embodiment in L2 have primarily focused on the lexical and phrasal levels, often providing conflicting and indeterminate results. Therefore, to address this gap, the present eye-tracking study aimed to explore the embodied mental representations formed during the reading comprehension of L2 action sentences. Specifically, it sought to identify the types of motor representations formed during L2 action sentence comprehension and the extent to which these representations are motor embodied.</p>
</sec>
<sec>
<title>Methods</title>
<p>A total of 56 advanced L2 learners participated in a Sentence&#x02013;Picture Verification Task, during which their response times (RTs) and eye movements were recorded. Each sentence&#x02013;picture pair depicted an action that either matched or mismatched the action implied by the sentence. Data analysis focused on areas of interest around the body effectors.</p>
</sec>
<sec>
<title>Results and discussion</title>
<p>RTs in the mismatch condition indicated an impeding effect. Furthermore, fixations on the body effector executing an action were longer in the mismatch condition, especially in late eye-movement measures.</p>
</sec></abstract>
<kwd-group>
<kwd>embodied cognition</kwd>
<kwd>eye tracking</kwd>
<kwd>L2 reading comprehension</kwd>
<kwd>mental representation</kwd>
<kwd>action sentence</kwd>
<kwd>motor</kwd>
<kwd>perception</kwd>
<kwd>ESL</kwd>
</kwd-group>
<counts>
<fig-count count="9"/>
<table-count count="2"/>
<equation-count count="0"/>
<ref-count count="59"/>
<page-count count="18"/>
<word-count count="12767"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Speech and Language</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="s1">
<title>1 Introduction</title>
<p>Amodal theories of language processing suggest that linguistic meaning is mentally represented by abstract symbols (Burgess and Lund, <xref ref-type="bibr" rid="B5">1997</xref>; Griffiths et al., <xref ref-type="bibr" rid="B22">2007</xref>). However, these theories cannot explain the perceptual experiences evoked when reading or hearing sentences (e.g., &#x0201C;a blue live lobster&#x0201D; vs. a &#x0201C;red cooked lobster&#x0201D;: Mimi looked at the lobster <italic>in the cold water</italic>/<italic>on a hot grill</italic>) (Glenberg et al., <xref ref-type="bibr" rid="B17">2004</xref>; Kiefer and Pulverm&#x000FC;ller, <xref ref-type="bibr" rid="B30">2012</xref>). Similarly, analyses based on amodal propositional representations overlook the differences in actions described by sentences like &#x0201C;Joyce is playing piano/violin&#x0201D; (i.e., key-pressing vs. bowing). To address these shortcomings, researchers who investigated embodied cognition have suggested that language comprehension involves the mental activation and integration of motor and perceptual experiences related to the events described by linguistic symbols (Zwaan, <xref ref-type="bibr" rid="B54">2004</xref>; Borghi, <xref ref-type="bibr" rid="B4">2012</xref>; Shapiro, <xref ref-type="bibr" rid="B43">2019</xref>).</p>
<p>Converging empirical evidence from neuroscience and behavioral research supports the theory that language meaning is grounded in our motor&#x02013;perceptual experiences of the world, thus favoring that language comprehension is embodied (e.g., Kaschak et al., <xref ref-type="bibr" rid="B29">2005</xref>; Tettamanti et al., <xref ref-type="bibr" rid="B47">2005</xref>; Casteel, <xref ref-type="bibr" rid="B6">2011</xref>; Moreno et al., <xref ref-type="bibr" rid="B35">2013</xref>; De Koning et al., <xref ref-type="bibr" rid="B10">2017</xref>; Kronrod and Ackerman, <xref ref-type="bibr" rid="B31">2021</xref>). One line of empirical evidence for this embodied approach to language meaning is action-based (Glenberg and Kaschak, <xref ref-type="bibr" rid="B18">2002</xref>), which emphasizes motor embodiment. At the lexical level, Hauk et al. (<xref ref-type="bibr" rid="B23">2004</xref>) demonstrated that L1 action verbs develop their motor representations somatotopically in the motor and premotor cortex. This has inspired studies in L2 embodied pedagogy studies, leading to the design of action-oriented classroom activities to facilitate L2 vocabulary learning (e.g., Ulbricht, <xref ref-type="bibr" rid="B49">2020</xref>; Garc&#x000ED;a-G&#x000E1;mez et al., <xref ref-type="bibr" rid="B16">2021</xref>), though not as much in L2 reading comprehension.</p>
<p>However, several L1 studies on embodied representations at the sentential level show that motor embodiment is context-dependent (e.g., Santana and De Vega, <xref ref-type="bibr" rid="B41">2013</xref>; Schuil et al., <xref ref-type="bibr" rid="B42">2013</xref>). A recent L2 study found that the motor representations of L2 action verbs are less or not grounded in motor experiences compared to L1 action verbs (Tian et al., <xref ref-type="bibr" rid="B48">2020</xref>). To the best of the authors&#x00027; knowledge, no L2 study of this kind has yet explored motor embodiment at the sentential level. Therefore, this study aimed to investigate which motor representations are constructed during the comprehension of L2 action sentences.</p>
<sec>
<title>1.1 Mental representations and embodiment</title>
<p>To understand a text, individuals not only represent the text itself but also build mental representations, or situation models, of the scenarios described (Zwaan and Radvansky, <xref ref-type="bibr" rid="B58">1998</xref>; Magliano et al., <xref ref-type="bibr" rid="B34">1999</xref>). An embodied account of language comprehension argues that motor&#x02013;perceptual representations are integral to constructing these mental representations (Barsalou, <xref ref-type="bibr" rid="B3">1999</xref>; Zwaan, <xref ref-type="bibr" rid="B53">1999</xref>). By activating and integrating motor&#x02013;perceptual experiences, individuals can gather and elaborate on information about the events described in a text, leading to a high level of language comprehension (Zwaan, <xref ref-type="bibr" rid="B55">2014</xref>; De Vega, <xref ref-type="bibr" rid="B11">2015</xref>). That is, mental representations are multimodal, involving perceptual and motor representations, and are therefore embodied. Constructing mental representations can be challenging for less skilled language users, as it requires higher-order and deep-level processing (Zwaan and Taylor, <xref ref-type="bibr" rid="B59">2006</xref>). Empirical evidence shows that L2 users often find tasks requiring such processing inefficient and effortful (Francis and Guti&#x000E9;rrez, <xref ref-type="bibr" rid="B15">2012</xref>; Horiba, <xref ref-type="bibr" rid="B26">1996</xref>; P&#x000E9;rez et al., <xref ref-type="bibr" rid="B38">2019</xref>). Thus, building mental representations appears particularly challenging for L2 users. Notably, Zwaan and Brown (<xref ref-type="bibr" rid="B56">1996</xref>) found that English learners of French developed weak and reduced mental representations of French stories, making fewer inferences, while their mental representations of English stories were strong and comprehensive.</p>
</sec>
<sec>
<title>1.2 Embodiment in L1</title>
<p>The compatibility effect, frequently observed in studies using a Sentence&#x02013;Picture Verification Task (SPVT), offers behavioral evidence that individuals mentally create motor&#x02013;perceptual representations to comprehend sentences in L1. For example, in Zwaan and Pecher&#x00027;s (<xref ref-type="bibr" rid="B57">2012</xref>) research on perceptual representation, a facilitating effect (i.e., a shorter reaction time, RT) was observed when the shape/orientation of an object implied in a sentence matched what was seen in a picture. Specifically, RTs were shorter when participants saw an image of a bird with outstretched wings after reading &#x0201C;There is a bird in the sky.&#x0201D; Conversely, an impeding effect, characterized by a longer RT, was observed when there was a mismatch between the sentence and the picture, such as when the sentence was &#x0201C;There is a bird in the nest.&#x0201D; Similarly, Holt and Beilock (<xref ref-type="bibr" rid="B25">2006</xref>) employed an SPVT to show that native English speakers mentally construct motor representations when reading action sentences in L1. In summary, the facilitating and impeding effects observed in the SPVT indicate the formation of motor or perceptual representations. Ferstl et al. (<xref ref-type="bibr" rid="B13">2017</xref>) investigated whether manipulating action performers&#x00027; facial appearances and clothing affected viewers&#x00027; judgments of differences and similarities between two ambiguous actions. Their results showed that visual representations of actions, along with the facial identities of the action performers, influenced viewers&#x00027; ability to distinguish between the actions. Notably, both Ferstl et al. (<xref ref-type="bibr" rid="B13">2017</xref>) and Holt and Beilock (<xref ref-type="bibr" rid="B25">2006</xref>) used picture stimuli depicting only the action performers&#x00027; hands in motion (i.e., representing affordances) without showing the actual objects involved. This similarity provides useful guidelines for displaying picture stimuli in motor embodiment studies.</p>
</sec>
<sec>
<title>1.3 Motor embodiment in L1: action-based sentence comprehension</title>
<p>The Indexical Hypothesis (IH) is an embodied approach to language comprehension. It emphasizes situated action and suggests that the meaning of a situation is determined by the series of actions individuals take to interact with the physical world. Accordingly, sentence comprehension relies on bodily action and involves three processes: indexing, extracting and inferring affordances, and meshing (Glenberg and Robertson, <xref ref-type="bibr" rid="B19">1999</xref>; Kaschak and Glenberg, <xref ref-type="bibr" rid="B28">2000</xref>). To determine whether a sentence is meaningful, individuals index or map words/phrases to their physical referents at the beginning of the process of understanding the sentence. Subsequently, they derive potential affordances related to these references. Affordances represent how individuals interact with referents and their applications in various situations. For example, to move luggage, travelers pull the collapsible handle. Finally, the grammatical structure of the sentence guides the meshing of these affordances to execute the intended action. Here is an example of how this process works: &#x0201C;He hangs the jacket on the collapsible handle of upright luggage.&#x0201D; The syntactic structure of this sentence guides readers to integrate specific actions implied in the sentence to achieve the goal of hanging up the jacket, such as extending the collapsible handle from the wheeled bag of the luggage and draping the jacket over it. However, meaninglessness arises when affordances are meshed into impossible actions (e.g., &#x0201C;he hangs the jacket on the upright bowl&#x0201D;). The ability to successfully mesh affordances into a coherent set of actions depends on the relationship between the affordances of the individual (i.e., he), the objects involved (i.e., jacket, luggage/bowl), and the goal specified in the sentence context.</p>
<p>Santana and De Vega (<xref ref-type="bibr" rid="B41">2013</xref>) illustrated the action-based sentence comprehension model. Their study found an enhanced N400 effect when participants read sentences describing a protagonist performing two actions simultaneously (e.g., <italic>while applying ointment to her wounded hand, she stretched a</italic> <italic><bold>bandage</bold></italic>). N400 is a brain signature sensitive to semantic or world knowledge inconsistencies. However, Santana and De Vega (<xref ref-type="bibr" rid="B41">2013</xref>) explained that the increased N400 observed in their study demonstrated a motor incongruency effect rather than a semantic violation, as their sentence stimuli did not contain semantic anomalies. Moreover, the two actions (i.e., stretching a bandage and applying ointment) were consistent with the scene described in the sentences, indicating that world knowledge violations did not contribute to the increased N400. This explanation was further supported by another experiment from the same study, which found that the N400 effect diminished when the two actions were performed consecutively. They explained further that the observed N400 reflected the unsuccessful formation of motor representations in the brain&#x00027;s motor regions due to a failure in the affordance meshing of the two actions. They also noted that the N400 effect peaked at the end of the sentence (i.e., bandage), not at the mention of the second verb (i.e., stretched). Based on these findings, Santana and de Vega suggested that motor representations are formed when sufficient contextual information is provided. This suggestion appears to conflict with the theory that action verbs alone activate motor embodiment (Hauk et al., <xref ref-type="bibr" rid="B23">2004</xref>). However, Schuil et al. (<xref ref-type="bibr" rid="B42">2013</xref>) found that the motor region was activated when action verbs were embedded in literal sentences (e.g., Peter picks up the books after the examination) but less activated in nonliteral sentences (e.g., Peter picks up the pieces after the examination). This highlights the essential role of sentence context in activating and integrating motor experiential traces for developing motor representations.</p>
</sec>
<sec>
<title>1.4 Embodied aspects of mental representations revealed through eye-tracking data</title>
<p>Oculomotor recording has proven to be a valuable tool for uncovering the mental representations formed for sentence comprehension (Anderson and Spivey, <xref ref-type="bibr" rid="B2">2009</xref>) and has been employed to explore embodied language comprehension. For example, dwell time and fixation counts on the depicted protagonist and destination were shorter and fewer for sentences describing fast actions compared to those describing slow actions (e.g., <italic>A man dashed/sauntered into the supermarket</italic>). In other words, the speed of a motion described in a sentence influenced the amount of gaze directed at the protagonist (i.e., the man) and the destination (i.e., the supermarket) relevant to the motion (Lindsay et al., <xref ref-type="bibr" rid="B33">2013</xref>; Speed and Vigliocco, <xref ref-type="bibr" rid="B45">2014</xref>). The changes in dwell time and fixation counts on the depicted protagonist and the destination indicate that indexing occurs in two directions. One approach maps words or phrases to their mental/internal representations (Glenberg and Robertson, <xref ref-type="bibr" rid="B20">2000</xref>), while the other links external/physical elements to these mental/internal representations (Spivey and Richardson, <xref ref-type="bibr" rid="B46">2009</xref>). The eye-tracking studies above showed that embodied indexing occurs in two directions simultaneously. Specifically, verbal inputs create mental representations, which are then linked/indexed to the externally and pictorially presented protagonist, destination, and event.</p>
</sec>
<sec>
<title>1.5 Embodiment in L2</title>
<p>Pavlenko (<xref ref-type="bibr" rid="B37">2014</xref>) argued that L2 is disembodied because bilinguals are generally exposed to L2 in grammar-based, amodal symbol-oriented instructional contexts. Using an SPVT, Norman and Peleg (<xref ref-type="bibr" rid="B36">2022</xref>) found that Hebrew L2 English readers did not show signs of constructing perceptual representations when comprehending English. Similarly, Chen et al. (<xref ref-type="bibr" rid="B8">2020</xref>), employing a delayed SPVT, found no perceptual embodiment in L2 Chinese and L3 English. Conversely, Vukovic and Williams (<xref ref-type="bibr" rid="B51">2014</xref>) showed that perceptual, experiential traces were automatically activated when English learners of Dutch processed English sentences describing perceptual information during an SPVT. Similarly, Ahn and Jiang (<xref ref-type="bibr" rid="B1">2018</xref>), using an SPVT, found that late L2 Korean learners developed perceptual representations for Korean sentences related to shape or orientation. Of particular interest is Foroni (<xref ref-type="bibr" rid="B14">2015</xref>) study, which observed that Dutch L2 English learners&#x00027; cheek muscles for smiling were stimulated when reading sentences expressing affirmative emotions (e.g., I am grinning). This finding was consistent with his previous L1 study of this kind. However, unlike L1 users in his previous study, L2 readers did not show inhibited cheek muscle activity when reading sentences expressing negative emotions (e.g., I am not grinning). Therefore, Foroni (<xref ref-type="bibr" rid="B14">2015</xref>) concluded that emotional language comprehension in L2 is less embodied than in L1 (see Zhang and Vanek, <xref ref-type="bibr" rid="B52">2021</xref> for a discussion of the processing cost of negation in L2). Regarding motor representations in L2 processing, Vukovic and Shtyrov (<xref ref-type="bibr" rid="B50">2014</xref>) observed that the motor experiential traces activated by action verbs were weaker for L2 (English) than for L1 (German). However, Tian et al. (<xref ref-type="bibr" rid="B48">2020</xref>) found the opposite. In their study, three types of English and Chinese verb phrases were produced as stimuli, with each conveying a different meaning: abstract, literal, and metaphorical. They found that motor activation strength increased from abstract to metaphorical to literal verb phrases in both languages. However, the motor activation for verb phrases was stronger in L2 (English) than in L1 (Chinese). They concluded that the strong motor responses observed in L2 verb phrases indicated that participants were processing less automatic language rather than solely reflecting motor representations. Therefore, they supported the view that L2 action-related language is processed in a disembodied manner.</p>
</sec>
<sec>
<title>1.6 The present study</title>
<p>Research on the role of perceptual traces in L2 comprehension has yielded mixed results. While previous L2 studies on perceptual embodiment have focused on lexical and sentential levels, studies on motor embodiment in L2 have been limited to the lexical and phrasal levels, often reporting conflicting and indeterminate results. However, L1 studies (Santana and De Vega, <xref ref-type="bibr" rid="B41">2013</xref>; Schuil et al., <xref ref-type="bibr" rid="B42">2013</xref>) have suggested that sentence context is crucial for activating and integrating motor experiential traces. Moreover, these traces are essential for forming motor representations needed to comprehend action sentences. Therefore, it is important to explore whether motor representations are developed at the sentential level in L2, consequently motivating this study to explore the motor aspects of mental representations (i.e., motor representations) formed for L2 action sentence comprehension. This study addresses the following research questions (RQs):</p>
<list list-type="simple">
<list-item><p>RQ 1: What embodied aspects of mental representations are involved in L2 action sentence comprehension? Specifically, what motor experiential traces contribute to these mental representations?</p></list-item>
<list-item><p>RQ 2: To what extent are these L2 mental representations grounded in motor experiential traces?</p></list-item>
</list>
<p>Using the SPVT paradigm and the eye-tracking technique, this study compared L2 readers&#x00027; eye gaze while viewing matched and mismatched sentence-picture pairs. Specifically, the picture depicted a protagonist performing an action that either corresponded with or differed from the action implied in the sentence. Readers&#x00027; eye gaze was recorded as they assessed whether the picture matched their comprehension of the sentence stimuli. As mentioned in Section 1.4, fixations index the development of mental representations. Therefore, analyzing participants&#x00027; eye gaze on the picture will reveal the motor representations they form to comprehend the sentence.</p>
<p>This study operationalized the formation of motor representations using IH. Specifically, the answers to RQ 1 were derived from analyzing how L2 participants gaze at and examine the action and the protagonist&#x00027;s body effectors (i.e., head, hands, and legs) in the picture stimuli. Because action performers&#x00027; identity (i.e., facial appearance/facial identity) affects observers&#x00027; ability to correctly identity the actions of the action performers (Ferstl et al., <xref ref-type="bibr" rid="B13">2017</xref>), our first hypothesis (HP) is as follows:</p>
<list list-type="simple">
<list-item><p>HP1</p></list-item>
<list-item><p><italic>If L2 participants form motor representations for comprehending an action sentence stimulus, their eye gaze will focus on the protagonist&#x00027;s head in the picture stimulus</italic>.</p></list-item>
</list>
<p>Specifically, participants were expected to index/map the protagonist&#x00027;s head in the picture to the protagonist described in the action sentence. That is, they would index the visual representation of the protagonist to the mental referent of the protagonist formed for comprehending the action sentence. This indexing would allow participants to determine if the depicted protagonist matched the one described in the sentence.</p>
<p>We also examined the relationships between the depicted body effectors, the implied action, and the described object based on the reference status in the sentence of the body effectors. The following hypotheses address affordances and meshing.</p>
<list list-type="simple">
<list-item><p>HP2</p></list-item>
<list-item><p><italic>If the affordances of the protagonist and the described object are successfully meshed, the eye gaze on the picture will shift to the body effector performing the action implied in the sentence. Moreover, the number of gazes on this body effector will be higher in the mismatch condition than in the match condition</italic>.</p></list-item>
<list-item><p>HP3</p></list-item>
<list-item><p><italic>When motor representations are developed, both the match&#x02013;mismatch conditions of the sentences and the reference status of the body effectors in the sentence will jointly influence participants&#x00027; eye gaze on the picture stimuli</italic>.</p></list-item>
</list>
<p>In summary, RQ 1 can be answered by examining which parts of the picture stimulus participants focus on and inspect and how frequently they do so. Although this study primarily focused on motor representations of hand actions, the picture stimuli also included the protagonist&#x00027;s legs. It was anticipated that the legs, compared with the hands, would have weaker relationships with the hand action and the described object. Consequently, participants were expected to focus more on the hands than on the legs. When the leg was presented in the picture stimuli, it was used as a distractor to assess whether the participants correctly formed the motor representation of hand actions. Previous studies employing the SPVT paradigm suggest that the facilitating and impeding effects reflect that motor&#x02013;perceptual representations were formed to comprehend the sentence stimuli before the picture stimuli were displayed (Holt and Beilock, <xref ref-type="bibr" rid="B25">2006</xref>). Therefore, we expected to see the two effects in this study.</p>
<p>Additionally, this study examined the extent to which the L2 mental representations are experientially embodied in motor terms, particularly focusing on the timing effect (i.e., early/late eye movements). Specifically, it investigated how quickly or slowly L2 participants identified actions depicted in picture stimuli that either matched or mismatched the actions described in the sentence stimuli. L2 mental representations tend to be reduced and weaker because of the slow and demanding nature of higher-order and deep-level processing for L2 users (Horiba, <xref ref-type="bibr" rid="B26">1996</xref>; Zwaan and Brown, <xref ref-type="bibr" rid="B56">1996</xref>). Based on previous findings, we propose the following hypothesis:</p>
<list list-type="simple">
<list-item><p>HP4:</p></list-item>
<list-item><p><italic>If L2 motor representations are weak and lack sufficient motor experiential traces, identifying whether the action depicted in the picture matches or mismatches the action described in the sentence will be slow and delayed. Conversely, stronger L2 motor representations will lead to quicker identification</italic>.</p></list-item>
</list>
</sec>
</sec>
<sec sec-type="materials and methods" id="s2">
<title>2 Materials and methods</title>
<sec>
<title>2.1 Participants</title>
<p>Overall, 56 right-handed L1 Chinese-speaking students (21 men, 35 women; mean age = 20.4 years, <italic>SD</italic> = 1.1) with normal or corrected-to-normal vision were recruited from universities in northern Taiwan. To participate, they were required to have taken a reading test from a globally recognized English proficiency test (e.g., TOEFL or IELTS) within the past 2 years and provide a copy of their test score report.</p>
<p>All participants had English reading abilities at the C1 level on the CEFR scale<xref ref-type="fn" rid="fn0001"><sup>1</sup></xref>. We initially recruited 79 students for the experiment, but 23 of them were excluded from further analyses due to the following reasons: minoring in English, high rates of incorrect responses on verification tasks and comprehension tests, invalid RTs, invalid fixation durations, and technical issues during data collection. Details on data cleaning and analysis are provided in Section 2.7. The remaining 56 participants were majoring in engineering or medicine.</p>
<p>G<sup>&#x0002A;</sup>Power software Version 3.1.9.7 (Faul et al., <xref ref-type="bibr" rid="B12">2007</xref>) was used to conduct a priori power analysis to estimate the optimal sample size for the 3 (body effectors) &#x000D7; 2 (match&#x02013;mismatch conditions) repeated measure design. A large effect size (<italic>f</italic><sup>2</sup>) of 0.35 was used, with an alpha of 0.05, a power of 0.80, and five predictor variables. The results showed that a minimum sample size of 43 was needed to detect a large effect size with an actual power of 0.81. Therefore, our final sample size (<italic>N</italic> = 56) was adequate. Superpower 0.20 (Lakens and Caldwell, <xref ref-type="bibr" rid="B32">2021</xref>) was used to conduct a simulation-based priori power analysis. We specified a sample size of 20 per cell, a common <italic>SD</italic> of 1, a common correlation of 0.70, and means of 1, 1.3, 1, 2.05, 0, and 0.1 for the six cells, respectively, with an alpha of 0.05. The results showed that this study would achieve an expected power of 0.90 with an effect size of <italic>f</italic><sup>2</sup> = 0.34. Thus, our final sample size of 28 per cell was adequate.</p>
</sec>
<sec>
<title>2.2 Experimental design</title>
<p>This study employed the SPVT combined with eye-tracking techniques to investigate how the reference status of body effectors and the match&#x02013;mismatch conditions in sentences influenced participants&#x00027; eye gaze on the three body effectors depicted in the picture stimuli. The study featured two repeated measures as independent variables: body effectors and match&#x02013;mismatch conditions. For body effectors, there were three levels: head, hand, and leg. Each body effector was categorized according to its reference status within the sentence, reflecting the relationships among the three body effectors, the implied action, and the described objects. For example, in the sentence &#x0201C;Mr. Bean is measuring the length of cars in a car park,&#x0201D; the name &#x0201C;Mr. Bean&#x0201D; explicitly identifies the protagonist and, by extension, the Head body effector. The Hand body effector, however, is implicitly involved, as the action &#x0201C;measuring&#x0201D; suggests the use of hands, though it is not directly mentioned in the sentence. The leg body effector, neither explicitly nor implicitly referenced, serves as a distractor. Therefore, the body effectors were categorized as follows:</p>
<list list-type="simple">
<list-item><p>Mentioned (Head, Hand) vs. Not-mentioned (Leg)</p></list-item>
<list-item><p>Explicitly mentioned (Head) vs. Implicitly mentioned (Hand)</p></list-item>
</list>
<p>For match&#x02013;mismatch conditions, we manipulated whether the action implied in the sentence matched or mismatched the action depicted in the picture (i.e., action-matching sentence vs. action-mismatching sentence).</p>
<p>To answer RQ 1, a 3 (Body effectors: Head, Hand, Leg) &#x000D7; 2 (Conditions: Match, Mismatch) repeated measure design was applied to eye movement data: fixation counts<xref ref-type="fn" rid="fn0002"><sup>2</sup></xref>. The independent variables were body effectors and match-mismatch conditions, with fixation counts as the dependent variable. Moreover, two 2 &#x000D7; 2 repeated measure designs were used for further analysis:</p>
<list list-type="simple">
<list-item><p>(A) 2 (Mentioned: Head and Hand vs. Not-mentioned: Leg) &#x000D7; 2 (Match vs. Mismatch Conditions).</p></list-item>
<list-item><p>(B) 2 (Explicitly mentioned: Head vs. Implicitly mentioned: Hand) &#x000D7; 2 (Match vs. Mismatch conditions).</p></list-item>
</list>
<p>Comprised two repeated measures of independent variables: mentioned vs. not-mentioned body effectors and match vs. mismatch conditions, each with two levels. Analysis of (A) determined whether and how these variables influenced participants&#x00027; fixation counts on the body effectors depicted in the picture stimuli. Similarly, (B) involved two repeated measures of independent variables, each with two levels. Analysis of (B) assessed how these variables influenced participants&#x00027; fixation counts on the body effectors depicted in the picture stimuli. To address RQ 2, a one-way repeated measures design was used, with match&#x02013;mismatch conditions serving as the independent variable and the eye movement data for temporal measures<xref ref-type="fn" rid="fn0003"><sup>3</sup></xref> in the Hand region serving as the dependent variable.</p>
</sec>
<sec>
<title>2.3 Materials</title>
<p>A total of 24 critical sentence&#x02013;picture item sets were created, each consisting of two sentences and one full-color picture drawn by a professional artist<xref ref-type="fn" rid="fn0004"><sup>4</sup></xref>. Each of the two sentences implied a different action (i.e., each item set had two implied actions). Both sentences contained the same verb, and the implied actions were performed by the same body effector of the same protagonist. However, the picture in each item set matched the action implied by only one of the sentences (i.e., action-matching sentence vs. action-mismatching sentence; see <xref ref-type="table" rid="T1">Table 1</xref>, <xref ref-type="fig" rid="F1">Figure 1</xref> for examples). During each trial, participants were shown only one sentence from a given critical sentence&#x02013;picture item set.</p>
<table-wrap position="float" id="T1">
<label>Table 1</label>
<caption><p>Some samples of SPVT pairs with comprehension test items.</p></caption>
<table frame="box" rules="all">
<tbody>
<tr style="background-color:#dee1e1">
<td valign="top" align="left" colspan="2"><bold>Sentence-picture pairs of a critical trial: sample 1</bold></td>
</tr>
<tr>
<td valign="top" align="left"><italic>The sentences in Sample 1 Mismatch</italic><break/> A boy is opening a can of Coke in the metro car.<break/> <italic>Match</italic><break/> A boy is opening a jar of jam in the metro car.<break/> <italic>Comprehension test item</italic>:<break/> A boy is opening a carton of milk in the metro car.</td>
<td valign="top" align="left"><xref ref-type="fig" rid="F1">Figure 1</xref>. The picture in sample 1.<break/> <inline-graphic xlink:href="fnhum-18-1410242-i0001.tif"/></td>
</tr>
<tr style="background-color:#dee1e1">
<td valign="top" align="left" colspan="2"><bold>Sentence-picture pairs of a critical trial: sample 2</bold></td>
</tr>
<tr>
<td valign="top" align="left"><italic>The Sentences in Sample 2 Match</italic><break/> A grandma is filling teacups with green tea in the dining room.<break/> <italic>Mismatch</italic><break/> A grandma is filling cupcakes with vanilla cream in the dining room. <italic>Comprehension test item</italic>:<break/> The grandma is filling a pie with minced beef in the dining room.</td>
<td valign="top" align="left"><xref ref-type="fig" rid="F4">Figure 4</xref>. The picture in sample 2.<break/> <inline-graphic xlink:href="fnhum-18-1410242-i0002.tif"/></td>
</tr>
<tr style="background-color:#dee1e1">
<td valign="top" align="left" colspan="2"><bold>A sample of filler pair</bold></td>
</tr>
<tr>
<td valign="top" align="left"><italic>The sentence in the sample of filler pair</italic><break/> A director is filming a hyena wrestling with its partner for food.<break/> <italic>Comprehension test item</italic>:<break/> A director is filming a hyena wrestling with its partner for territory.</td>
<td valign="top" align="left"><xref ref-type="fig" rid="F2">Figure 2</xref>. The picture in the sample of the filler pair.<break/> <inline-graphic xlink:href="fnhum-18-1410242-i0003.tif"/></td>
</tr></tbody>
</table>
</table-wrap>
<fig id="F1" position="float">
<label>Figure 1</label>
<caption><p>Picture in sample 1.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0001.tif"/>
</fig>
<p>The verbs were selected based on their definitions in the Cambridge and Oxford English Dictionaries, with the criterion that neither dictionary could explicitly indicate which body effector executes the action described by the verb. For example, the verb &#x0201C;catch&#x0201D; was excluded because its definition, &#x0201C;to stop and hold a moving object or person, especially in your hands,&#x0201D; implies the use of hands.</p>
<p>This study focused on making inferences about hand- and arm-related actions. While the way an action is performed and the body effector associated with the action were not explicitly stated, they could be inferred from the context. Each sentence presented to the participants explicitly mentioned a protagonist. Sentence length and the protagonist&#x00027;s position in the sentence were consistent across sentence stimuli in both the match and mismatch conditions. In summary, although both sentences in any given sentence&#x02013;picture item set contained the same verb, the action depicted in the picture matched the action in only one of the sentences. Notably, the specific details of the action were not explicitly mentioned but needed to be inferred from the sentence context. Additionally, filler pairs were included to prevent participants from discerning the purpose of the study. These filler sentences described objects, locations, or general situations and were paired with pictures that were similar but irrelevant to the meaning of the sentences (<xref ref-type="table" rid="T1">Table 1</xref>, <xref ref-type="fig" rid="F2">Figure 2</xref>). Moreover, yes/no comprehension tests were developed for critical and filler items to verify that participants understood the sentences used in the verification tasks. For instance, the comprehension test for the &#x0201C;boy&#x0201D; sentence&#x02013;picture item set shown in <xref ref-type="table" rid="T1">Table 1</xref> is: &#x0201C;A boy is opening a carton of milk in the metro car.&#x0201D;</p>
<fig id="F2" position="float">
<label>Figure 2</label>
<caption><p>Picture in the sample of filler pair.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0002.tif"/>
</fig>
<p>A separate group of 10 advanced L2 readers participated in a pilot study. The recruitment criteria of these participants were similar to those for the main study. The pilot study aimed to confirm that none of the words in the sentences&#x02014;whether critical items or fillers&#x02014;were unfamiliar to the participants and that the correct action could be inferred from each sentence stimulus. In the pilot study, participants rated the compatibility of each sentence stimulus with the corresponding illustrated action on a 7-point Likert scale, ranging from &#x0201C;very compatible&#x0201D; to &#x0201C;very incompatible.&#x0201D; Sentence&#x02013;picture pair rated below 5 was excluded. Additionally, items that consistently induced prolonged RTs or unusual eye gaze<xref ref-type="fn" rid="fn0005"><sup>5</sup></xref> were excluded. Ultimately, 12 critical item sets and 24 fillers were retained for the experiment. The average compatibility rating of the critical items in the match condition was 6.5 (<italic>SD</italic> = 0.60).</p>
<p>In summary, two lists of SPVT item sets (List 1 and List 2) were created. Both lists contained the same filler items. The critical items and sentence conditions (match and mismatch) were counterbalanced across the two lists. This means that if a critical item set presented its match condition on one list, it showed its mismatch condition on the other. Specifically, each picture in a critical item set appeared once in each list, paired with the action-matching sentence on one list and the action-mismatching sentence on the other. Moreover, each list comprised one practice block and three experimental blocks, with an equal number of filler pairs and matching and mismatching pairs in each block. Within each block, the critical pairs and filler pairs were randomly displayed. The order of the three experimental blocks was also randomized. Comprehension tests were administered after each block. Participants were randomly assigned to one of the two lists: half viewed List 1, while the other half viewed List 2.</p>
</sec>
<sec>
<title>2.4 Procedure</title>
<p>Participants began with practice trials to familiarize themselves with the task. A 9-point calibration was used to adjust the eye-tracking apparatus. During the experiment, drift correction was performed before each sentence&#x02013;picture pair presentation. Each sentence was displayed on the screen one at a time, centered and left-justified. A black cross was displayed for 100 ms before each sentence stimulus to direct participants&#x00027; attention. The black cross was positioned to the left of the center of the screen, on the first letter of each sentence. Once participants had read and understood the sentence, they pressed the space bar. A black cross then appeared in the center of the screen for 100 ms, followed by a blank screen for 250 ms. A picture stimulus was then displayed either to the left or right of the center of the screen. The second black cross was intended to focus participants&#x00027; visual attention before presenting the picture stimuli.</p>
<p>Participants were instructed to quickly determine whether the picture depicted any elements (e.g., a protagonist, place, object) mentioned in the preceding sentence, and their RTs were recorded. They used the computer keyboard to respond, pressing the &#x0201C;P&#x0201D; key labeled &#x0201C;Yes&#x0201D; if they believed the image was mentioned in the sentence and the &#x0201C;Q&#x0201D; key for &#x0201C;No&#x0201D; if they did not. Participants were not explicitly required to determine whether or not the action depicted in the picture matched the action implied in the sentence. Instead, they were simply asked to respond to &#x0201C;whether any element in the picture had been mentioned in the immediately preceding sentence.&#x0201D; The expected response for all of the critical trials was &#x0201C;yes,&#x0201D; while a &#x0201C;no&#x0201D; response was expected for all the fillers. As noted in Section 2.3, participants were randomly assigned to one of the two SPVT item set lists. Each list included one practice block and three experimental blocks. The comprehension test items were presented at the end of each block. <xref ref-type="fig" rid="F3">Figure 3</xref> illustrates the SPVT flowchart.</p>
<fig id="F3" position="float">
<label>Figure 3</label>
<caption><p>SPVT flowchart.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0003.tif"/>
</fig>
</sec>
<sec>
<title>2.5 Apparatus</title>
<p>The EyeLink 1000 Plus eye-tracking system (SR Research Ltd.) with a 1,000 Hz sampling rate was used to record eye movements. Only the dominant eye was monitored, although viewing was binocular. Pictures and sentences were displayed on a 21-inch LCD monitor with a 1,440 &#x000D7; 900 screen resolution and a 50 Hz refresh rate. Participants sat 69 cm from the screen, with their heads stabilized by a chin-and-forehead rest to minimize head movement. The picture stimuli subtended a visual angle of 33.27&#x000B0; &#x000D7; 39.85&#x000B0; (width &#x000D7; height). Finally, sentences were displayed in 24-point Courier New font, black on a white background.</p>
</sec>
<sec>
<title>2.6 Measurements</title>
<p>This study used a standard SPVT approach, measuring the RT required for participants to determine if anything presented in the picture was mentioned in the immediately preceding sentence. This method helped avoid inadvertently revealing the research objective to the participants during the task. Only RTs for the critical items were included in the statistical analysis. Additionally, eye movement data were analyzed to assess the processing time as participants visually attended to predetermined areas of interest (AOIs) in the picture stimuli.</p>
<p>The measurements are detailed in <xref ref-type="table" rid="T2">Table 2</xref>. Each body effector (the Hands, Head, and Legs) was assigned a specific AOI. Additionally, an All-Inclusive AOI covered the protagonist in the image (the whole picture) (<xref ref-type="fig" rid="F4">Figure 4</xref>). The four AOIs, shown in the picture of the older woman in <xref ref-type="table" rid="T1">Table 1</xref>, enabled both general and specific inferences of the data. Each of the six eye movement measures detailed in <xref ref-type="table" rid="T2">Table 2</xref> corresponds to a different time event (Godfroid, <xref ref-type="bibr" rid="B21">2019</xref>). The Hand AOI was used to assess data on indicators (a)&#x02013;(f), the All-Inclusive AOI was used for (e) and (f), and the Head and Leg AOIs for (f). This study focused exclusively on motor representations of hand- and arm-related actions. Therefore, the results for RQ 2, indicated by temporal eye movement measures, reflect the degree of motor embodiment of the Hand action developed by the participants. Conversely, RQ 1 results, shown by fixation counts, revealed which parts of the picture participants examined. Overall, RQ 1 was addressed through fixation counts (f) in the Head, Hand, and Leg AOIs, while RQ 2 was answered through temporal measures (a)&#x02013;(e) in the Hand AOI.</p>
<table-wrap position="float" id="T2">
<label>Table 2</label>
<caption><p>Eye movement measures adopted in this study.</p></caption>
<table frame="box" rules="all">
<thead>
<tr style="background-color:#919498;color:#ffffff">
<th valign="top" align="left"><bold>Measure</bold></th>
<th valign="top" align="left"><bold>Definition</bold></th>
<th valign="top" align="left"><bold>Indicator of Timing Effects</bold></th>
</tr>
</thead>
<tbody>
<tr>
<td valign="top" align="left">(a) First fixation duration</td>
<td valign="top" align="left">The time spent fixating on an AOI for the first time, reflects the processes of identification and recognition (De Graef et al., <xref ref-type="bibr" rid="B9">1990</xref>)</td>
<td valign="top" align="left">Initial processing time (i.e., early eye movement measures)</td>
</tr>
<tr>
<td valign="top" align="left">(b) First-pass dwell time</td>
<td valign="top" align="left">The sum of the duration of all fixations falling within an AOI before moving to another AOI, which reflects the process of object recognition and tends to increase when an unexpected object is noticed (Henderson et al., <xref ref-type="bibr" rid="B24">1999</xref>)</td>
<td valign="top" align="left">Early</td>
</tr>
<tr>
<td valign="top" align="left">(c) Second-pass dwell time</td>
<td valign="top" align="left">The time spent on the previously fixated AOI, excluding the first-pass dwell time, suggests that the participant engaged in reanalysis and intentional processing of the AOI</td>
<td valign="top" align="left">Late processing time (i.e., late eye movement measures)</td>
</tr>
<tr>
<td valign="top" align="left">(d) Regression-in-count</td>
<td valign="top" align="left">The number of times that an AOI, which is currently being viewed after having been processed previously, is returned to while looking at other AOIs, or in other words, the number of revisits</td>
<td valign="top" align="left">Late</td>
</tr>
<tr>
<td valign="top" align="left">(e) Total dwell time</td>
<td valign="top" align="left">The sum of all the dwell times for the same AOI during a single trial is an indication of the overall cognitive effort made while processing stimuli</td>
<td valign="top" align="left">Late</td>
</tr>
<tr>
<td valign="top" align="left">(f) Fixation count (i.e., the number of fixations falling within an AOI)</td>
<td valign="top" align="left">The strength of the participants&#x00027; attentiveness to the AOI</td>
<td valign="top" align="left">Global indicator</td>
</tr></tbody>
</table>
</table-wrap>
<fig id="F4" position="float">
<label>Figure 4</label>
<caption><p>Picture in sample 2.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0004.tif"/>
</fig>
<p>The accuracy of the responses to the reading comprehension tests and the verification task were recorded. The reading comprehension test was designed to elicit a simple yes/no response (<xref ref-type="table" rid="T1">Table 1</xref>). The test results reflected the participants&#x00027; attention to and understanding of each sentence stimulus. The comprehension tests were included based on the idea that people form comprehensive mental representations of what they read only when they read for understanding or learning (Zwaan, <xref ref-type="bibr" rid="B55">2014</xref>). Accordingly, data were removed if the accuracy of either the reading comprehension tests or the SPVT was lower than 80%<xref ref-type="fn" rid="fn0006"><sup>6</sup></xref>.</p>
</sec>
<sec>
<title>2.7 Statistical analyses and data cleaning</title>
<p>This study used a repeated measures design. All data analyses were conducted using R software, Version 3.6.1 (R Core Team, <xref ref-type="bibr" rid="B40">2019</xref>). Criteria for valid fixation and RT were strictly enforced for data cleaning. If any fixation or RT from a participant did not meet these criteria, all data collected from that participant were removed from the analyses<xref ref-type="fn" rid="fn0007"><sup>7</sup></xref>. The data are accessible at this website: <ext-link ext-link-type="uri" xlink:href="https://osf.io/zeufn/s">https://osf.io/zeufn/s</ext-link>.</p>
<sec>
<title>2.7.1 Behavioral data cleaning and analysis</title>
<p>The average accuracy rates for the verification task and comprehension test were 96.3% and 87%, respectively. Only RTs within three <italic>SDs</italic> of a participant&#x00027;s mean RT (mean &#x000B1; 3 <italic>SD</italic>) for any single trial were considered valid. To determine whether impeding and facilitating effects occurred, a pairwise t-test was performed on the RTs.</p>
</sec>
<sec>
<title>2.7.2 Eye-tracking data cleaning and analysis</title>
<p>The temporal threshold for identifying valid ocular fixation was set at 50 ms<xref ref-type="fn" rid="fn0008"><sup>8</sup></xref>. Despite strict data cleaning, the second-pass dwell time data remained asymmetrically distributed. This issue was resolved by applying a square-root transformation, resulting in a better approximation of a normal distribution.</p>
<p>To answer RQ 1, we first computed a pairwise <italic>t</italic>-test on total dwell time and fixation counts to determine if global eye movement measures reflect impeding and facilitating effects related to motor embodiment. We then employed a 3 (Body effectors: Head, Hand, Leg) &#x000D7; 2 (Conditions: Match, Mismatch) repeated measures design to test the hypotheses of RQ 1. Data analyses were conducted on fixation counts using a linear mixed-effects model with the <italic>lme</italic> function in the <italic>nlme</italic> package, Version 3.1-141 (Pinheiro et al., <xref ref-type="bibr" rid="B39">2021</xref>). To analyze the main effects and interactions, three orthogonal contrasts were set. Contrast 1 compared the match-to-mismatch condition of the sentences. Contrasts 2 and 3 compared the relationships among the three body effectors, the implied action, and the described object (the reference status in the sentence). Specifically, Contrasts 2 and 3 explored whether the participants&#x00027; inspection of the picture stimuli was impacted by the mention of the body effector and whether it was explicit or implicit. In summary, Contrast 1 compared match vs. mismatch conditions, Contrast 2 compared mentioned vs. not-mentioned body effectors, and Contrast 3 compared explicitly vs. implicitly mentioned body effectors. Moreover, four models were built to identify the best fit for the fixation count data. The commands to run the four models in R are accessible at <ext-link ext-link-type="uri" xlink:href="https://osf.io/zeufn/s">https://osf.io/zeufn/s</ext-link>. The specifications of the four models are as follows:</p>
<p><italic>The baseline model</italic> included only an intercept, predicting the outcome solely from the intercept. The random part of this model showed that the two repeated-measures predictors, body effectors, and match-mismatch conditions, were nested within the participant variable. Moreover, maximum likelihood estimation was used for this model. Next, one predictor was added at a time. Keeping the same outcomes and predictors as <italic>the baseline model</italic>, we added body effectors to <italic>the baseline model</italic> as a predictor using the <italic>update()</italic> function to build <italic>the body effector model</italic>. Then, match&#x02013;mismatch conditions were added to <italic>the body effectors model</italic> to create <italic>the match&#x02013;mismatch conditions model</italic>. Finally, the interaction between match&#x02013;mismatch conditions and body effectors was added to the <italic>match&#x02013;mismatch conditions model</italic> to create <italic>the interaction model</italic>.</p>
<p>Every model was estimated using maximum likelihood, considering the random effects of the participants. The four models were compared using the <italic>anova()</italic> function. <italic>Post-hoc</italic> tests for the simple main effect of mentioned vs. not-mentioned body effectors, match vs. mismatch conditions, and explicitly vs. implicitly mentioned body effectors were conducted using pairwise <italic>t</italic>-tests with Bonferroni correction. Notably, the mentioned body effectors included fixation counts on Head AOI and Hand AOI. The simple main effect of the match vs. mismatch conditions in the 2 (mentioned vs. not-mentioned) &#x000D7; 2 (match vs. mismatch conditions) interaction was computed to compare the fixation counts on the not-mentioned Leg with the averaged fixation counts on the mentioned Head and Hand [(Head &#x0002B; Hand)/2 vs. Leg]. Afterward, a simple main effect of body effectors was assessed using the <italic>glht</italic> function in the <italic>multcomp</italic> package, Version 1.4-10 (Hothorn et al., <xref ref-type="bibr" rid="B27">2008</xref>). Finally, paired <italic>t</italic>-tests with Bonferroni correction were performed on the selected eye-movement temporal measures to answer RQ 2.</p>
</sec>
</sec>
</sec>
<sec sec-type="results" id="s3">
<title>3 Results</title>
<sec>
<title>3.1 RTs during SPVT: facilitating and impeding effects were observed</title>
<p>The pairwise t-test was conducted to determine whether match-mismatch conditions affected RTs, which required participants to judge compatibility after viewing the picture stimulus. On average, RTs in the mismatch condition (<italic>M</italic> = 2,179.91, <italic>SD</italic> = 1,021.51) were significantly longer than in the match condition (<italic>M</italic> = 1,360.08, <italic>SD</italic> = 528.46, <italic>t</italic>(335) = &#x02212;12.98, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 0.98, 95% CI [&#x02212;944.04, &#x02212;695.62]). This finding indicates a facilitating effect in the match condition and an impeding effect in the mismatch condition.</p>
</sec>
<sec>
<title>3.2 Whole-picture AOI: eye movement data reflect facilitating and impeding effects</title>
<p>When the entire picture was set as a single AOI (whole-picture AOI), the pairwise <italic>t</italic>-test showed that this region received more fixations in the mismatch condition (<italic>M</italic> = 5.02, <italic>SD</italic> = 2.02; <italic>M</italic> = 3.82, <italic>SD</italic> = 1.74) than in the match condition (<italic>M</italic> = 3.82, <italic>SD</italic> = 1.74; <italic>t</italic>(335) = &#x02212;9.33, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 0.60, 95% CI [&#x02212;1.08, &#x02212;1.67]). Similarly, the paired samples <italic>t</italic>-test also demonstrated that total viewing time was significantly longer in the mismatch condition (<italic>M</italic> = 1,802.92, <italic>SD</italic> = 981.08) than in the match condition (<italic>M</italic> = 1,061.07, <italic>SD</italic> = 528.18; <italic>t</italic>(335) = &#x02212;12.13, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 0.90, 95% CI [&#x02212;862.10, &#x02212;621.59]). This indicates an impeding effect in the mismatch condition.</p>
</sec>
<sec>
<title>3.3 Results for RQ 1: what embodied aspects of mental representations are formed to achieve L2 action sentence comprehension? Specifically, what motor experiential traces are involved in these mental representations? Three (Body effectors: Head, Hand, Leg) &#x000D7; Two (Conditions: Match, Mismatch) repeated measure design</title>
<p>RQ 1 explored which parts of the picture stimulus participants looked at and inspected, and how often they looked at those parts. Fixation count data collected for RQ 1 were analyzed using a linear mixed model approach. The contrasts built for the analyses were match vs. mismatch conditions, mentioned vs. not-mentioned body effectors, and explicitly vs. implicitly mentioned body effectors. Four models were also constructed: <italic>the baseline model, body effectors model, match-mismatch conditions model, and interaction model</italic> (see Sections 2.2 and 2.7 for details on contrasts and models). To identify the best-fitting model, the four models were compared using the <italic>anova()</italic> function. The results showed that <italic>the body effectors model</italic> fit the data better than <italic>the baseline model</italic> [&#x003C7;<sup>2</sup>(2) = 482.81, <italic>p</italic> &#x0003C; 0.001], and <italic>the match&#x02013;mismatch conditions model</italic> fit better than <italic>the body effectors model</italic> [&#x003C7;<sup>2</sup>(1) = 53.45, <italic>p</italic> &#x0003C; 0.001]. However, the best fit was the <italic>interaction model</italic> [&#x003C7;<sup>2</sup>(2) = 79.05, <italic>p</italic> &#x0003C; 0.001], as confirmed by the Akaike Information Criterion: <italic>the baseline model</italic> = 5,747, <italic>the body effectors model</italic> = 5,268, <italic>the match&#x02013;mismatch conditions model</italic> = 5,217, and <italic>the interaction model</italic> = 5,142. This indicates that the fixation counts were significantly affected by the match-mismatch conditions and the type of body effectors (Match: Hand: <italic>M</italic> = 1.74, <italic>SD</italic> = 0.89; Head: <italic>M</italic> = 1.81, <italic>SD</italic> = 0.97; Leg: <italic>M</italic> = 0.06, <italic>SD</italic> = 0.27; Mismatch: Hand: <italic>M</italic> = 2.82, <italic>SD</italic> = 1.19; Head: <italic>M</italic> = 2.09, <italic>SD</italic> = 1.07; Leg: <italic>M</italic> = 0.15, <italic>SD</italic> = 0.44). <xref ref-type="fig" rid="F5">Figure 5</xref> shows the interaction graph, illustrating the inspection pattern concerning fixation counts across the three body effectors, focusing on the non-parallel lines. The line representing the match condition illustrates that fixation counts were slightly higher on the Head than on the Hand. However, the line representing the mismatch condition showed a different trend: fixation counts on the Hand were higher than on the Head. This match and mismatch conditions seemed not to affect the Leg. The higher position of the mismatch condition compared to the match condition line indicated that the mismatch condition resulted in higher fixation counts on the Hand and the Head. Further analyses of the interaction are provided in the following sections.</p>
<fig id="F5" position="float">
<label>Figure 5</label>
<caption><p>Main interaction effect.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0005.tif"/>
</fig>
<sec>
<title>3.3.1 Two (Mentioned: Head and Hand vs. Not-mentioned: Leg) &#x000D7; Two (Match vs. Mismatch conditions) interaction</title>
<p>The first interaction term examined the effect of mentioned body effectors (Head and Hand) relative to the not-mentioned body effector (Leg) in match and mismatch conditions (Match: <italic>M</italic> = 3.55, <italic>SD</italic> = 1.37; Mismatch: <italic>M</italic> = 4.90, <italic>SD</italic> =1.72; After average, Head &#x0002B; Hand/2; Match: <italic>M</italic> = 1.77, <italic>SD</italic> = 6.85; Mismatch: <italic>M</italic> = 2.45, <italic>SD</italic> = 8.60). The contrast was significant (<xref ref-type="fig" rid="F6">Figure 6</xref>). The effect of the mentioned body effectors compared to the not-mentioned body effector in increasing inspection eye gaze (fixation counts) was significantly stronger in the mismatch condition than in the match condition [<italic>b</italic> = &#x02212;0.19, <italic>t</italic>(220) = &#x02212;6.28, <italic>p</italic> &#x0003C; 0.001, <italic>f</italic><sup>2</sup> = 0.20]. Moreover, the main effect of mentioned vs. not-mentioned body effector was significant (<italic>b</italic> = 1.33, <italic>SE</italic> = 0.03<italic>, t</italic> = 42.92<italic>, p</italic>&#x0003C;<italic>0.0</italic>01, 95% CI = [1.27, 1.39]), as was the main effect of match-mismatch conditions (<italic>b</italic> = &#x02212;0.23, <italic>SE</italic> = 0.03<italic>, t</italic> = &#x02212;8.26<italic>, p</italic> &#x0003C; 0.001, 95% CI = [&#x02212;0.29, &#x02212;0.18]). The simple main effect of match vs. mismatch conditions showed that the mentioned Head &#x0002B; Hand had significantly higher fixation counts in the mismatch condition than in the match condition (<italic>t</italic>(335) = 11.30, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 0.87, 95% CI [1.11, 1.58]). Although differences in fixation counts were also found between the match and mismatch conditions for the not-mentioned Leg, the effect size was minimal (d &#x0003C; 0.2); therefore, these differences could be ignored (<italic>t</italic>(335) = &#x02212;2.97, <italic>p</italic> = 0.05, <italic>d</italic> = 0.17, 95% CI [&#x02212;0.02, &#x02212;0.14]). Conversely, the simple main effect of mentioned vs. not-mentioned body effectors showed that fixation counts on the mentioned Head &#x0002B; Hand were significantly higher than on the not-mentioned Leg in both the match (<italic>t</italic>(335) = 44.33, <italic>p</italic> = &#x0003C; 0.001, <italic>d</italic> = 0.35, 95% CI [1.63, 1.78]) and mismatch conditions (<italic>t</italic>(335) = 46.23, <italic>p</italic> = &#x0003C; 0.001, <italic>d</italic> = 0.38, 95% CI [2.20, 2.39]).</p>
<fig id="F6" position="float">
<label>Figure 6</label>
<caption><p>Interaction between mentioned vs. not-mentioned body effector and match vs. mismatch conditions.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0006.tif"/>
</fig>
</sec>
<sec>
<title>3.3.2 Two (Explicitly mentioned: Head vs. Implicitly mentioned: Hand) &#x000D7; Two (Match vs. Mismatch conditions) interaction</title>
<p>The second interaction term examined whether the effect of the explicitly mentioned body effector compared to the implicitly mentioned body effector differed between the match and mismatch conditions. The contrast was significant (<xref ref-type="fig" rid="F7">Figure 7</xref>). The effect of the implicitly mentioned body effector (Hand) compared to the explicitly mentioned body effector (Head) on fixation counts was significantly stronger in the mismatch condition than in the match condition (<italic>b</italic> = 0.20, <italic>t</italic>(220) = 7.41, <italic>p</italic> &#x0003C; 0.001, <italic>f</italic><sup>2</sup> = 0.30). The main effect of explicitly vs. implicitly mentioned body effectors was also significant (<italic>b</italic> = &#x02212;0.16, <italic>SE</italic> = 0.03<italic>, t</italic> = &#x02212;0.63<italic>, p</italic> &#x0003C; 0.001, 95% CI = [&#x02212;0.21, &#x02212;0.10]). The simple main effect of match vs. mismatch conditions showed that fixation counts differed between the match and mismatch conditions for the explicitly mentioned Head (<italic>t</italic>(335) = &#x02212;3.41, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 0.19, 95% CI [&#x02212;0.43, &#x02212;0.11]), though the effect size was minimal and the differences were negligible. Conversely, the implicitly mentioned Hand received significantly higher fixations in the mismatch condition than in the match condition (<italic>t</italic>(335) = &#x02212;13.24, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 1.09, 95% CI [&#x02212;1.23, &#x02212;0.91]). The simple main effect of explicitly vs. implicitly mentioned body effectors revealed that, in the match condition, fixation counts on the explicitly mentioned Head and the implicitly mentioned Hand were comparable (<italic>t</italic>(335) = 1.06, <italic>p</italic> = 0.28, 95% CI [&#x02212;0.06, 0.21]). Finally, in the mismatch condition, fixation counts on the implicitly mentioned Hand were significantly higher than those on the explicitly mentioned Head (<italic>t</italic>(335) = 8.96, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 0.65, 95% CI [0.56, 0.88]).</p>
<fig id="F7" position="float">
<label>Figure 7</label>
<caption><p>Interactions between explicitly vs. implicitly mentioned body effectors and match vs. mismatch conditions.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0007.tif"/>
</fig>
</sec>
<sec>
<title>3.3.3 Simple main effects of body effectors</title>
<p>We first built a baseline model with only intercept and random effects for participants, estimated using maximum likelihood to analyze the main effect of the body effectors. We then built two additional models: <italic>the match condition model</italic>, which included the match condition as a predictor, and <italic>the mismatch condition model</italic>, which included the mismatch condition as a predictor. The likelihood ratio was computed to compare <italic>the baseline model</italic> with <italic>the match condition model</italic> and <italic>the baseline model</italic> with <italic>the mismatch condition model</italic> [Match: <inline-formula><mml:math id="M1"><mml:msubsup><mml:mrow><mml:mi>&#x003C7;</mml:mi></mml:mrow><mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mn>2</mml:mn></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:msubsup></mml:math></inline-formula> = 244.09, <italic>p</italic> &#x0003C; 0.001; Mismatch: <inline-formula><mml:math id="M2"><mml:msubsup><mml:mrow><mml:mi>&#x003C7;</mml:mi></mml:mrow><mml:mrow><mml:mrow><mml:mo stretchy="false">(</mml:mo><mml:mrow><mml:mn>2</mml:mn></mml:mrow><mml:mo stretchy="false">)</mml:mo></mml:mrow></mml:mrow><mml:mrow><mml:mn>2</mml:mn></mml:mrow></mml:msubsup></mml:math></inline-formula> = 343.75, <italic>p</italic> &#x0003C; 0.001]. Finally, <italic>post-hoc</italic> tests showed that, in the match condition, fixation counts on the Hand and the Head were similar (<italic>z</italic> = 0.94, <italic>p</italic> = 0.61, 95% CI [0.25, &#x02212;0.10]). Conversely, fixation counts on the Leg in the match condition were significantly fewer than those on the Hand (<italic>z</italic> = &#x02212;21.35, <italic>p</italic> &#x0003C; 0.001, <italic>f</italic><sup>2</sup>= 1.28, 95% CI [&#x02212;1.85, &#x02212;1.49]) and the Head (<italic>z</italic> = &#x02212;22.30, <italic>p</italic> &#x0003C; 0.001, <italic>f</italic><sup>2</sup> = 1.5, 95% CI [&#x02212;1.93, &#x02212;1.56]). In the mismatch condition, fixation counts on the Head were significantly fewer than those on the Hand (<italic>z</italic> = &#x02212;9.82, <italic>p</italic> &#x0003C; 0.001, <italic>f</italic><sup>2</sup> = 0.28, 95% CI [&#x02212;0.89, &#x02212;0.55]). Fixation counts on the Leg were also significantly fewer than those on the Hand (<italic>z</italic> = &#x02212;36.03, <italic>p</italic> &#x0003C; 0.001, <italic>f</italic><sup>2</sup> = 3.76, 95% CI [&#x02212;2.83, &#x02212;2.49]) and the Head (<italic>z</italic> = &#x02212;26.20, <italic>p</italic> &#x0003C; 0.001, <italic>f</italic><sup>2</sup>= 0.81, 95% CI [&#x02212;2.11, &#x02212;1.76]).</p>
<p>The results indicate that the not-mentioned body effector (Leg) did not capture participants&#x00027; attention as much as the explicitly and implicitly mentioned body effectors (Head and Hands). This trend is illustrated in the count-based fixation and duration-based fixation maps presented in <xref ref-type="fig" rid="F8">Figures 8</xref>, <xref ref-type="fig" rid="F9">9</xref>. The color legend to the right of the image explains the mapping: warmer colors, such as red, represent higher fixation counts and longer durations.</p>
<fig id="F8" position="float">
<label>Figure 8</label>
<caption><p>Count-based fixation map of total sample in match (<bold>left</bold>) and mismatch (<bold>right</bold>) conditions.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0008.tif"/>
</fig>
<fig id="F9" position="float">
<label>Figure 9</label>
<caption><p>Duration-based fixation map of total sample in match (<bold>left</bold>) and mismatch (<bold>right</bold>) conditions.</p></caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fnhum-18-1410242-g0009.tif"/>
</fig>
</sec>
</sec>
<sec>
<title>3.4 Results for RQ 2: to what extent are L2 mental representations grounded in motor experiential traces? Specifically, how fast did the participants identify whether the action depicted in the picture matches or mismatches the action implied in the sentence?</title>
<p>To answer RQ 2, we conducted a series of paired t-tests. The analysis revealed no differences between the match and mismatch conditions in the first fixation duration on the Hand (match: <italic>M</italic> = 240.71, <italic>SD</italic> = 151.15; mismatch: <italic>M</italic> = 255.27, <italic>SD</italic> = 117.54, <italic>t</italic>(335) = &#x02212;1.59, <italic>p</italic> = 0.11, 95% CI [&#x02212;32.50, 3.39]). Conversely, the first-pass dwell time on the Hand was notably shorter in the match condition than in the mismatch condition (match: <italic>M</italic> = 293.54, <italic>SD</italic> = 157.54; mismatch: <italic>M</italic> = 395.97, <italic>SD</italic> = 198.39, <italic>t</italic>(335) = &#x02212;7.51, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 0.50, 95% CI [&#x02212;129.24, &#x02212;75.60]).</p>
<p>When comparing total dwell time on the Hand between the match and mismatch conditions, it was observed that participants fixated on the Hand significantly longer in the mismatch condition than in the match condition (Match: <italic>M</italic> = 400.51, <italic>SD</italic> = 192.43; Mismatch: <italic>M</italic> = 633.76, <italic>SD</italic> = 310.41, <italic>t</italic>(335) = &#x02212;11.52, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 0.88, 95% CI [&#x02212;273.04, &#x02212;193.44]). Additionally, a square-root transformation of the second-pass dwell time data showed that dwell time was significantly shorter in the match condition than in the mismatch condition (match: <italic>M</italic> = 7.00, <italic>SD</italic> = 7.62; Mismatch: <italic>M</italic> = 14.24, <italic>SD</italic> = 5.92, <italic>t</italic>(335) = &#x02212;13.28, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 1.15, 95% CI [&#x02212;8.35, &#x02212;6.19]). Finally, there was a significant difference between the match and mismatch conditions for the regression-in-count measure. The participants revisited the Hand more frequently in the mismatch condition than in the match condition (match: <italic>M</italic> = 0.50, <italic>SD</italic> = 0.50; Mismatch: <italic>M</italic> = 1.29, <italic>SD</italic> = 0.71, <italic>t</italic>(335) = &#x02212;16.71, <italic>p</italic> &#x0003C; 0.001, <italic>d</italic> = 1.49, 95% CI [&#x02212;0.89, &#x02212;0.70]). Ultimately, this indicates a significant impeding effect in the mismatch condition, primarily evident in late temporal eye movement measures.</p>
</sec>
</sec>
<sec sec-type="discussion" id="s4">
<title>4 Discussion</title>
<p>This study proposed two research questions to explore motor embodiment in L2 comprehension. The first research question investigated which embodied aspects of mental representations were formed to achieve L2 action sentence comprehension. The second question probed the extent to which these mental representations were grounded in motor experiential traces. Two major findings emerged. First, in the mismatch condition, participants&#x00027; fixation counts were highest for the implicitly mentioned Hand, followed by the explicitly mentioned Head, and lowest for the Leg. Second, the temporal threshold for identifying the mismatch action was at the first-pass-dwell time. The following discusses how the results were interpreted to answer the two research questions.</p>
<sec>
<title>4.1 Embodied aspects of mental representations are formed for L2 action sentence comprehension</title>
<p>Previous studies on L1 embodiment using the SPVT paradigm have shown that RT facilitation and impediment in the match and mismatch conditions, respectively, indicate that motor&#x02013;perceptual representations were mentally formed for sentence comprehension (Zwaan and Pecher, <xref ref-type="bibr" rid="B57">2012</xref>; De Koning et al., <xref ref-type="bibr" rid="B10">2017</xref>). Similarly, the present study on L2 using the SPVT found both facilitating and impeding effects. Notably, these effects were also found in global eye movement measures, such as total dwell time on the whole-picture AOI. This indicates that L2 readers form embodied aspects of mental representations for L2 action sentence comprehension before the picture stimuli are presented, as evidenced by their eye gaze when inspecting the picture stimuli.</p>
</sec>
<sec>
<title>4.2 Eye gaze on picture stimuli revealed motor aspects of mental representations</title>
<p>RQ 1 investigated whether and what motor experiential traces are involved in mental representations formed for L2 action sentence comprehension. To answer this, we examined which parts of the picture stimulus participants focused on and inspected. HP 3 proposed that the match&#x02013;mismatch conditions of the sentences and the reference status in the sentence of the body effectors would jointly influence participants&#x00027; eye gaze on the picture stimuli. The results supported HP 3, as the <italic>interaction model</italic> best fits the data on fixation counts. The findings related to fixation counts on each body effector, along with HP 1 and HP 2 for RQ1, are discussed in the following sections.</p>
<sec>
<title>4.2.1 HP 1: if L2 participants form motor representations for comprehending an action sentence stimulus, their eye gaze will focus on the protagonist&#x00027;s head in the picture stimulus</title>
<p>The mean fixation count on the Leg in match and mismatch conditions is close to zero (&#x0003C; 0.2), indicating that the Leg bears little relation to the implied action and the affordance of the described objects. This finding is expected since the sentence stimuli require inferences related to the hand and arm, and the Leg was neither explicitly nor implicitly mentioned. Conversely, the mean fixation counts on the Hand and Head were above 1 in both conditions. This indicates that the Hand and Head were consistently inspected. Embodied indexing fixation emerges when external objects are mapped to their mental referents (Spivey and Richardson, <xref ref-type="bibr" rid="B46">2009</xref>). Therefore, the participants&#x00027; focus on the Head in the picture in match and mismatch conditions signals that they were aligning the protagonist in the picture with the mental referent of the protagonist described in the sentence.</p>
<p>HP 1 was confirmed by comparing the mean fixation counts on the Head and Hand. The mean fixation counts on the Head and Hand were comparable in the match condition, indicating that the Head and Hand are equally relevant to the implied action and the affordance of the described objects. On the other hand, the mean fixation counts on the Head in the mismatch condition were comparable with those in the match condition but were significantly fewer than those on the Hand in the mismatch condition. This suggests that the Head was not the body effector executing the action implied in the sentence. However, it played a role in identifying whether the action depicted in the picture matched or mismatched the action implied in the sentence. In this regard, our findings exemplify Ferstl et al.&#x00027;s (<xref ref-type="bibr" rid="B13">2017</xref>) observation that recognizing an action performed by a protagonist is impacted by the protagonist&#x00027;s facial identity.</p>
</sec>
<sec>
<title>4.2.2 HP2: if the affordances of the protagonist and the described object are successfully meshed, the eye gaze on the picture will fall on the body effector executing the action implied in the sentence. Moreover, the number of inspections on this body effector will be higher in the mismatch condition than in the match condition</title>
<p>HP2 was supported by the finding that mean fixation counts on the Hand in the mismatch condition were significantly higher than those in the match condition and on the Head in the mismatch condition. This finding is notable because the Head, which signifies the protagonist&#x00027;s facial identity (Ferstl et al., <xref ref-type="bibr" rid="B13">2017</xref>), was explicitly mentioned in the sentence stimuli, whereas the Hand was not. Despite this, fixation counts on the Hand exceeded those on the Head in the mismatch condition. This indicates that the Hand was crucial for determining whether the depicted action matched or mismatched the action implied in the sentence. Moreover, these findings align with Glenberg and Robertson&#x00027;s (<xref ref-type="bibr" rid="B19">1999</xref>) IH, indicating that affordances were correctly extracted, inferred, and then successfully meshed. The successful meshing of affordances facilitated the development of motor representations and was crucial for determining whether the action shown in the picture matched or mismatched and could or could not be mapped onto the verbally induced and internally presented motor representations. Notably, the mean fixation counts on the Hand in the mismatch condition were significantly higher than those in the match condition. This implies that participants identified the inconsistency between the depicted action in the picture and the motor representations they had formed based on the implied action.</p>
<p>In summary, comparing the mean fixation counts of the three body effectors revealed that participants formed motor representations of hand actions. Consistent with Glenberg and Robertson (<xref ref-type="bibr" rid="B19">1999</xref>), these fixation counts indicated that motor representations were formed by the relationships among body effectors, actions, and the objects affected by those actions. These findings aligned with embodied indexing observed in previous oculomotor studies of L1 embodiment (Lindsay et al., <xref ref-type="bibr" rid="B33">2013</xref>; Speed and Vigliocco, <xref ref-type="bibr" rid="B45">2014</xref>), where language users mapped pictorial elements to the verbally induced mental representations. Importantly, the results of RQ 1 showed that L2 readers formed motor aspects of embodied mental representations to comprehend L2 action sentences. These findings, consistent with Zwaan and Pecher (<xref ref-type="bibr" rid="B57">2012</xref>) and Ahn and Jiang (<xref ref-type="bibr" rid="B1">2018</xref>), support the idea that L2 reading comprehension is embodied, similar to the results found in L1 studies conducted by Holt and Beilock (<xref ref-type="bibr" rid="B25">2006</xref>). Our study, similar to Holt and Beilock&#x00027;s (<xref ref-type="bibr" rid="B25">2006</xref>), investigates motor embodiment, but we focus on L2, whereas they focus on L1.</p>
</sec>
</sec>
<sec>
<title>4.3 The second earliest temporal indicator marked the temporal point at which the mismatched action was identified</title>
<p>RQ 2 explored to what extent these L2 mental representations are grounded in motor experiential traces. Specifically, HP 4 examined the speed at which participants determine whether the action depicted in a picture matches or mismatches the action implied in a sentence. The results revealed significant differences in fixation on the Hand between the match and mismatch conditions. These differences were found in both early temporal indicators (first-pass dwell time) and late temporal indicators (second-pass dwell time, regression-in-count, and total dwell time). Recall that first-pass dwell time indexes initial processing, with a longer duration indicating the detection of an improbable object (Henderson et al., <xref ref-type="bibr" rid="B24">1999</xref>). Therefore, the significant differences found on the first-pass dwell time indicated that during the first-pass dwell time, participants in the mismatch condition began to notice that the action depicted in the picture mismatched the action implied in the sentence. This raises the question of why these advanced L2 readers detected the mismatch at the second earliest temporal measure (first-pass dwell time) rather than the earliest (first fixation duration).</p>
<p>There are two reasons for this. First, as noted by De Graef et al. (<xref ref-type="bibr" rid="B9">1990</xref>), participants initially identified the object (Hand) in the picture, regardless of whether the condition matched or mismatched with the implied action. In other words, without first identifying the depicted limb or hand, participants would not have been able to determine whether the depicted action matched or mismatched the implied action.</p>
<p>Alternatively, the earliest time event rarely yields a significant effect in L2 studies compared to L1 due to processing speed differences between L1 and L2 (Godfroid, <xref ref-type="bibr" rid="B21">2019</xref>). L2 processing speed is generally slower and more arduous than L1. It is well-documented that L2 higher-order and deep-level processing is inefficient and effortful (Francis and Guti&#x000E9;rrez, <xref ref-type="bibr" rid="B15">2012</xref>; Horiba, <xref ref-type="bibr" rid="B26">1996</xref>; P&#x000E9;rez et al., <xref ref-type="bibr" rid="B38">2019</xref>). According to Zwaan and Brown (<xref ref-type="bibr" rid="B56">1996</xref>), mental representations in L2 are typically weaker and less detailed than those in L1, often due to fewer inferences being made. This supports the idea that participants in this study developed mental representations with a limited amount of motor experiential traces. In other words, their motor representations were less detailed and less clear regarding the actions implied in the sentences. The clarity of motor representations was sufficient for participants to recognize the action depicted in the picture in the match condition, but it was inadequate in the mismatch condition. This explains why all the late temporal eye-movement indicators, such as total dwell time and regression-in-count, on the Hand in this study, showed significant differences with large effect sizes between the match and mismatch conditions. Participants doubted whether the mismatched action in the picture truly differed from the action implied in the sentence due to the weak and unclear motor representations that they had formed. This uncertainty made them reinspect the conflicting part of the picture more frequently, increasing the regression-in-count and total dwell time on the Hand in the mismatch condition. Therefore, similar to the findings of Vukovic and Shtyrov (<xref ref-type="bibr" rid="B50">2014</xref>) and Foroni (<xref ref-type="bibr" rid="B14">2015</xref>), which suggest that L2 processing is less embodied in the perceptual aspect, this study indicates that L2 processing is also less embodied in the motor aspect.</p>
<p>Overall, this study shows that advanced L2 readers mentally form sufficient motor embodiment for L2 action sentence comprehension to identify mismatched actions, though not as quickly. Exploring how L2 proficiency differences impact temporal eye movement measures and the strength of motor representations is beyond the scope of this study.</p>
</sec>
<sec>
<title>4.4 Situated embodied L2 action language comprehension</title>
<p>This study shows that advanced L2 learners mentally form motor representations to comprehend action sentences. These representations include adequate motor experiential traces related to the action implied in the sentence, such as specific hand actions. These findings contradict the conclusions drawn by Pavlenko (<xref ref-type="bibr" rid="B37">2014</xref>) and Tian et al. (<xref ref-type="bibr" rid="B48">2020</xref>), who argued that L2 comprehension is disembodied. Although both Tian et al. (<xref ref-type="bibr" rid="B48">2020</xref>) and this study involved Chinese-speaking learners of English, Tian et al. employed lexical and phrasal stimuli. According to Borghi (<xref ref-type="bibr" rid="B4">2012</xref>) and Glenberg and Kaschak (<xref ref-type="bibr" rid="B18">2002</xref>), sentence context provides situational clues that guide the implementation of actions and goals. In other words, without sentence contexts, individuals struggle to mesh affordances. The conflicting results between Tian et al. and this study indicate that situational context is crucial for activating motor experiential traces, which promote motor representations formed for L2 action language comprehension. Thus, this study corroborates the findings of Santana and De Vega (<xref ref-type="bibr" rid="B41">2013</xref>) and Schuil et al. (<xref ref-type="bibr" rid="B42">2013</xref>), highlighting the importance of sentence context in activating motor experiential traces.</p>
</sec>
<sec>
<title>4.5 From amodal to embodied learning and teaching</title>
<p>Contrary to the amodal symbols (e.g., translation practices) and grammatical structures commonly emphasized in L2 classrooms, recent L2 embodied pedagogy studies focus on improving vocabulary learning through activities such as mime, gestures, and multimodal inputs (Ulbricht, <xref ref-type="bibr" rid="B49">2020</xref>; Garc&#x000ED;a-G&#x000E1;mez et al., <xref ref-type="bibr" rid="B16">2021</xref>). While these activities aim to promote bodily experiences of learning, they often overlook the importance of situated context in embodied language comprehension. According to the IH, action should occur within a situated context. Therefore, effective embodied language pedagogy should incorporate activities that engage learners in scenarios requiring specific actions, thereby promoting motor&#x02013;perceptual experiences relevant and meaningful within the learning context. Since adult L2 learners have already developed a range of situated embodied experiences through their L1, L2 embodied reading pedagogy should aim to connect and enhance these motor&#x02013;perceptual, experiential traces to L2. For instance, group discussions about visually representing text content can help adult EFL learners with low proficiency tap into their existing motor&#x02013;perceptual experiences (Shiang, <xref ref-type="bibr" rid="B44">2018</xref>).</p>
</sec>
</sec>
<sec id="s5">
<title>5 Conclusion and limitations</title>
<p>This study is the first to investigate the motor aspect of mental representations formed during the comprehension of L2 action sentences and to assess how these mental representations are embodied at the sentence level. Compared to previous L2 studies that focused on the lexical level, our findings advance the embodied account of language comprehension by demonstrating that L2 embodied action-language comprehension is context-driven. Overall, this study indicates that L2 is not disembodied. However, some critical item sets contain events with varying levels of daily familiarity, making direct comparisons challenging. For example, &#x0201C;opening a jar of jam&#x0201D; is not as common as &#x0201C;opening a can of coke.&#x0201D; The trade-off between event familiarity and the familiarity of words or phrases for EFL readers, as well as the clarity of the visual presentation of an action, was a factor in this study. Future research could explore how event familiarity impacts readers&#x00027; responses. Additionally, due to the large and complex nature of eye-tracking data analysis, this study primarily used data from a single task. Future research could investigate the generalizability of these findings to other tasks. This study only included advanced L2 learners; therefore, the results may not apply to L2 learners of other proficiency levels.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s6">
<title>Data availability statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec sec-type="ethics-statement" id="s7">
<title>Ethics statement</title>
<p>The studies involving human participants were reviewed and approved by the dissertation committee of the Department of English, National Taiwan Normal University. Informed consent was obtained from all individual participants included in the study.</p>
</sec>
<sec sec-type="author-contributions" id="s8">
<title>Author contributions</title>
<p>R-FS: Conceptualization, Data curation, Formal analysis, Investigation, Methodology, Project administration, Resources, Software, Validation, Visualization, Writing &#x02013; original draft, Writing &#x02013; review &#x00026; editing. C-LC: Supervision, Writing &#x02013; review &#x00026; editing. H-CC: Formal analysis, Methodology, Software, Writing &#x02013; original draft.</p>
</sec>
<sec sec-type="funding-information" id="s9">
<title>Funding</title>
<p>The author(s) declare that no financial support was received for the research, authorship, and/or publication of this article.</p>
</sec>
<ack><p>R-FS would like to thank her parents for their patience and support, which made completing this study possible. This study is dedicated to the memory of her father, Peter, Y-J, Shiang, a true gentleman and the best father in the world.</p>
</ack>
<sec sec-type="COI-statement" id="conf1">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s10">
<title>Publisher&#x00027;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<fn-group>
<fn id="fn0001"><p><sup>1</sup>Given that the action sentences in the SPVT were designed to induce mental representations of a text which involves high-level language processing, only proficient L2 readers were recruited.</p></fn>
<fn id="fn0002"><p><sup>2</sup>See Section 2.6 for an explanation.</p></fn>
<fn id="fn0003"><p><sup>3</sup>Also see Section 2.6 for an explanation.</p></fn>
<fn id="fn0004"><p><sup>4</sup>The artist also provided advice on which action can be clearly illustrated and presented in picture.</p></fn>
<fn id="fn0005"><p><sup>5</sup>The excluded items were those containing celebrities as protagonists, some participants inclined to look at the key features of these celebrities (e.g., Kim Kardashian).</p></fn>
<fn id="fn0006"><p><sup>6</sup>The threshold was based on previous studies on perceptual embodiment, such as Ahn and Jiang (<xref ref-type="bibr" rid="B1">2018</xref>) with 80%; Chen et al. (<xref ref-type="bibr" rid="B8">2020</xref>) with Chinese participants&#x00027; verification task at 83.05% and comprehension task 70.35%, and English participants&#x00027; verification task at 86.75% and comprehension task at 74.43%.</p></fn>
<fn id="fn0007"><p><sup>7</sup>As eye-tracking data for L2 learners is usually skewed (Godfroid, <xref ref-type="bibr" rid="B21">2019</xref>), a strict data cleaning approach was employed to minimize the risk of violating statistical assumptions and the complexity of the analyses.</p></fn>
<fn id="fn0008"><p><sup>8</sup>Although Castelhano and Henderson (<xref ref-type="bibr" rid="B7">2008</xref>) suggested that viewers can grasp the gist of a visual scene in about 40 ms, the thresholds for fixation count and for RT in this study were chosen based on our examination of the overall raw data.</p></fn>
</fn-group>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ahn</surname> <given-names>S.</given-names></name> <name><surname>Jiang</surname> <given-names>N.</given-names></name></person-group> (<year>2018</year>). <article-title>Automatic semantic integration during L2 sentential reading</article-title>. <source>Biling.</source> <volume>21</volume>, <fpage>375</fpage>&#x02013;<lpage>383</lpage>. <pub-id pub-id-type="doi">10.1017/S1366728917000256</pub-id></citation>
</ref>
<ref id="B2">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Anderson</surname> <given-names>S. E.</given-names></name> <name><surname>Spivey</surname> <given-names>M. J.</given-names></name></person-group> (<year>2009</year>). <article-title>The enactment of language: decades of interactions between linguistic and motor processes</article-title>. <source>Lang. Cogn</source>. <volume>1</volume>, <fpage>87</fpage>&#x02013;<lpage>111</lpage>. <pub-id pub-id-type="doi">10.1515/LANGCOG.2009.005</pub-id></citation>
</ref>
<ref id="B3">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Barsalou</surname> <given-names>L. W.</given-names></name></person-group> (<year>1999</year>). <article-title>Perceptual symbol systems</article-title>. <source>J. Behav. Brain Sci.</source> <volume>22</volume>, <fpage>577</fpage>&#x02013;<lpage>660</lpage>. <pub-id pub-id-type="doi">10.1017/S0140525X99002149</pub-id><pub-id pub-id-type="pmid">11301525</pub-id></citation></ref>
<ref id="B4">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Borghi</surname> <given-names>A. M.</given-names></name></person-group> (<year>2012</year>). <article-title>&#x0201C;Language comprehension: action, affordances and goals,&#x0201D;</article-title> in <source>Language and Action in Cognitive Neuroscience</source>, eds. Y. Coello and A. Bartolo (<publisher-loc>London</publisher-loc>: <publisher-name>Psychology Press</publisher-name>), <fpage>143</fpage>&#x02013;<lpage>162</lpage>.</citation>
</ref>
<ref id="B5">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Burgess</surname> <given-names>C.</given-names></name> <name><surname>Lund</surname> <given-names>K.</given-names></name></person-group> (<year>1997</year>). <article-title>Modelling parsing constraints with high-dimensional context space</article-title>. <source>Lang. Cogn. Process</source>. <volume>12</volume>, <fpage>177</fpage>&#x02013;<lpage>210</lpage>. <pub-id pub-id-type="doi">10.1080/016909697386844</pub-id></citation>
</ref>
<ref id="B6">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Casteel</surname> <given-names>M. A.</given-names></name></person-group> (<year>2011</year>). <article-title>The influence of motor simulations on language comprehension</article-title>. <source>Acta Psychol.</source> <volume>138</volume>, <fpage>211</fpage>&#x02013;<lpage>218</lpage>. <pub-id pub-id-type="doi">10.1016/j.actpsy.2011.06.006</pub-id><pub-id pub-id-type="pmid">21763635</pub-id></citation></ref>
<ref id="B7">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Castelhano</surname> <given-names>M. S.</given-names></name> <name><surname>Henderson</surname> <given-names>J. M.</given-names></name></person-group> (<year>2008</year>). <article-title>The influence of color on the perception of scene gist</article-title>. <source>J. Exp. Psychol: Hum. Percept. Perform</source>. <volume>34</volume>, <fpage>660</fpage>&#x02013;<lpage>675</lpage>. <pub-id pub-id-type="doi">10.1037/0096-1523.34.3.660</pub-id><pub-id pub-id-type="pmid">18505330</pub-id></citation></ref>
<ref id="B8">
<citation citation-type="journal"><person-group person-group-type="author"><collab>Chen D. Wang R. Zhang and, C. Liu</collab></person-group>. (<year>2020</year>). <article-title>Perceptual Representations in L1, L2 and L3 Comprehension: Delayed Sentence&#x02013;Picture Verification</article-title>. <source>J. Psycholinguist Res</source>. <volume>49</volume>, <fpage>41</fpage>&#x02013;<lpage>57</lpage>. <pub-id pub-id-type="doi">10.1007/s10936-019-09670-x</pub-id><pub-id pub-id-type="pmid">31468246</pub-id></citation></ref>
<ref id="B9">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>De Graef</surname> <given-names>P.</given-names></name> <name><surname>Christiaens</surname> <given-names>D.</given-names></name> <name><surname>d&#x00027;Ydewalle</surname> <given-names>G.</given-names></name></person-group> (<year>1990</year>). <article-title>Perceptual effects of scene context on object identification</article-title>. <source>Psychol. Res.</source> <volume>52</volume>, <fpage>317</fpage>&#x02013;<lpage>329</lpage>. <pub-id pub-id-type="doi">10.1007/BF00868064</pub-id><pub-id pub-id-type="pmid">2287695</pub-id></citation></ref>
<ref id="B10">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>De Koning</surname> <given-names>B. B.</given-names></name> <name><surname>Wassenburg</surname> <given-names>S. I.</given-names></name> <name><surname>Bos</surname> <given-names>L. T.</given-names></name> <name><surname>Van der Schoot</surname> <given-names>M.</given-names></name></person-group> (<year>2017</year>). <article-title>Mental simulation of four visual object properties: similarities and differences as assessed by the sentence&#x02013;picture verification task</article-title>. <source>J. Cogn. Psychol.</source> <volume>29</volume>, <fpage>420</fpage>&#x02013;<lpage>432</lpage>. <pub-id pub-id-type="doi">10.1080/20445911.2017.1281283</pub-id></citation>
</ref>
<ref id="B11">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>De Vega</surname> <given-names>M.</given-names></name></person-group> (<year>2015</year>). <article-title>&#x0201C;Toward an embodied approach to inferences in comprehension: The case of action language,&#x0201D;</article-title> in <source>Inferences during Reading</source>, eds. E. O&#x00027;Brien, A. Cook, and R. Lorch, Jr (Cambridge: Cambridge University Press), <fpage>182</fpage>&#x02013;<lpage>209</lpage>. <pub-id pub-id-type="doi">10.1017/CBO9781107279186.010</pub-id></citation>
</ref>
<ref id="B12">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Faul</surname> <given-names>F.</given-names></name> <name><surname>Erdfelder</surname> <given-names>E.</given-names></name> <name><surname>Lang</surname> <given-names>A. G.</given-names></name> <name><surname>Buchner</surname> <given-names>A.</given-names></name></person-group> (<year>2007</year>). <article-title>G<sup>&#x0002A;</sup>Power 3: a flexible statistical power analysis program for the social, behavioral, and biomedical sciences</article-title>. <source>Behav. Res. Methods</source> <volume>39</volume>, <fpage>175</fpage>&#x02013;<lpage>191</lpage>. <pub-id pub-id-type="doi">10.3758/BF03193146</pub-id><pub-id pub-id-type="pmid">17695343</pub-id></citation></ref>
<ref id="B13">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ferstl</surname> <given-names>Y.</given-names></name> <name><surname>B&#x000FC;lthoff</surname> <given-names>H.</given-names></name> <name><surname>De la Rosa</surname> <given-names>S.</given-names></name></person-group> (<year>2017</year>). <article-title>Action recognition is sensitive to the identity of the actor</article-title>. <source>Cognition</source> <volume>166</volume>, <fpage>201</fpage>&#x02013;<lpage>206</lpage>. <pub-id pub-id-type="doi">10.1016/j.cognition.2017.05.036</pub-id><pub-id pub-id-type="pmid">28582683</pub-id></citation></ref>
<ref id="B14">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Foroni</surname> <given-names>F.</given-names></name></person-group> (<year>2015</year>). <article-title>Do we embody second language? Evidence for &#x02018;partial&#x00027; simulation during processing of a second language</article-title>. <source>Brain Cogn.</source> <volume>99</volume>, <fpage>8</fpage>&#x02013;<lpage>16</lpage>. <pub-id pub-id-type="doi">10.1016/j.bandc.2015.06.006</pub-id><pub-id pub-id-type="pmid">26188846</pub-id></citation></ref>
<ref id="B15">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Francis</surname> <given-names>W. S.</given-names></name> <name><surname>Guti&#x000E9;rrez</surname> <given-names>M.</given-names></name></person-group> (<year>2012</year>). <article-title>Bilingual recognition memory: Stronger performance but weaker levels-of-processing effects in the less fluent language</article-title>. <source>Mem. Cognit</source>. <volume>40</volume>, <fpage>496</fpage>&#x02013;<lpage>503</lpage>. <pub-id pub-id-type="doi">10.3758/s13421-011-0163-3</pub-id><pub-id pub-id-type="pmid">22086650</pub-id></citation></ref>
<ref id="B16">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Garc&#x000ED;a-G&#x000E1;mez</surname> <given-names>A. B.</given-names></name> <name><surname>Cervilla</surname> <given-names>&#x000D3;.</given-names></name> <name><surname>Casado</surname> <given-names>A.</given-names></name> <name><surname>Macizo</surname> <given-names>P.</given-names></name></person-group> (<year>2021</year>). <article-title>Seeing or acting? The effect of performing gestures on foreign language vocabulary learning</article-title>. <source>Lang. Teach. Res.</source> 1055&#x02013;1086. <pub-id pub-id-type="doi">10.1177/13621688211024364</pub-id></citation>
</ref>
<ref id="B17">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Glenberg</surname> <given-names>A. M.</given-names></name> <name><surname>Gutierrez</surname> <given-names>T.</given-names></name> <name><surname>Levin</surname> <given-names>J. R.</given-names></name> <name><surname>Japuntich</surname> <given-names>S.</given-names></name> <name><surname>Kaschak</surname> <given-names>M. P.</given-names></name></person-group> (<year>2004</year>). <article-title>Activity and imagined activity can enhance young children&#x00027;s reading comprehension</article-title>. <source>J. Edu. Psychol</source>. <volume>96</volume>:<fpage>424</fpage>. <pub-id pub-id-type="doi">10.1037/0022-0663.96.3.424</pub-id></citation>
</ref>
<ref id="B18">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Glenberg</surname> <given-names>A. M.</given-names></name> <name><surname>Kaschak</surname> <given-names>M. P.</given-names></name></person-group> (<year>2002</year>). <article-title>Grounding language in action</article-title>. <source>Psychon. Bull. Rev</source>. <volume>9</volume>, <fpage>558</fpage>&#x02013;<lpage>565</lpage>. <pub-id pub-id-type="doi">10.3758/BF03196313</pub-id><pub-id pub-id-type="pmid">12412897</pub-id></citation></ref>
<ref id="B19">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Glenberg</surname> <given-names>A. M.</given-names></name> <name><surname>Robertson</surname> <given-names>D. A.</given-names></name></person-group> (<year>1999</year>). <article-title>Indexical understanding of instructions</article-title>. <source>Discourse Proc.</source> 28,1&#x02013;26. <pub-id pub-id-type="doi">10.1080/01638539909545067</pub-id></citation>
</ref>
<ref id="B20">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Glenberg</surname> <given-names>A. M.</given-names></name> <name><surname>Robertson</surname> <given-names>D. A.</given-names></name></person-group> (<year>2000</year>). <article-title>Symbol grounding and meaning: A comparison of high-dimensional and embodied theories of meaning</article-title>. <source>J. Mem. Lang.</source> <volume>43</volume>, <fpage>379</fpage>&#x02013;<lpage>401</lpage>. <pub-id pub-id-type="doi">10.1006/jmla.2000.2714</pub-id><pub-id pub-id-type="pmid">17950260</pub-id></citation></ref>
<ref id="B21">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Godfroid</surname> <given-names>A.</given-names></name></person-group> (<year>2019</year>). <source>Eye Tracking in Second Language Acquisition and Bilingualism: A Research Synthesis and Methodological Guide</source>. <publisher-loc>New York</publisher-loc>: <publisher-name>Routledge</publisher-name>.</citation>
</ref>
<ref id="B22">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Griffiths</surname> <given-names>T. L.</given-names></name> <name><surname>Steyvers</surname> <given-names>M.</given-names></name> <name><surname>Tenenbaum</surname> <given-names>J. B.</given-names></name></person-group> (<year>2007</year>). <article-title>Topics in semantic representation</article-title>. <source>Psychol. Rev.</source> <volume>114</volume>, <fpage>211</fpage>&#x02013;<lpage>244</lpage>. <pub-id pub-id-type="doi">10.1037/0033-295X.114.2.211</pub-id><pub-id pub-id-type="pmid">17500626</pub-id></citation></ref>
<ref id="B23">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hauk</surname> <given-names>O.</given-names></name> <name><surname>Johnsrude</surname> <given-names>I.</given-names></name> <name><surname>Pulverm&#x000FC;ller</surname> <given-names>F.</given-names></name></person-group> (<year>2004</year>). <article-title>Somatotopic representation of action words in human motor and premotor cortex</article-title>. <source>Neuron</source> <volume>41</volume>, <fpage>301</fpage>&#x02013;<lpage>307</lpage>. <pub-id pub-id-type="doi">10.1016/S0896-6273(03)00838-9</pub-id><pub-id pub-id-type="pmid">14741110</pub-id></citation></ref>
<ref id="B24">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Henderson</surname> <given-names>J. M.</given-names></name> <name><surname>Weeks Jr</surname> <given-names>P. A.</given-names></name> <name><surname>Hollingworth</surname> <given-names>A.</given-names></name></person-group> (<year>1999</year>). <article-title>The effects of semantic consistency on eye movements during complex scene viewing</article-title>. <source>J. Exp. Hum. Percept. Perform.</source> <volume>25</volume>, <fpage>210</fpage>&#x02013;<lpage>228</lpage>. <pub-id pub-id-type="doi">10.1037/0096-1523.25.1.210</pub-id></citation>
</ref>
<ref id="B25">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Holt</surname> <given-names>L. E.</given-names></name> <name><surname>Beilock</surname> <given-names>S. L.</given-names></name></person-group> (<year>2006</year>). Expertise and its embodiment: Examining the impact of sensorimotor skill expertise on the representation of action-related text., <italic>Psychon. Bull. Rev</italic>. <volume>13</volume>, <fpage>694</fpage>&#x02013;<lpage>701</lpage>. <pub-id pub-id-type="doi">10.3758/BF03193983</pub-id><pub-id pub-id-type="pmid">17201372</pub-id></citation></ref>
<ref id="B26">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Horiba</surname> <given-names>Y.</given-names></name></person-group> (<year>1996</year>). <article-title>Comprehension processes in L2 reading: Language competence, textual coherence, and inferences</article-title>. <source><italic>Stud. Second Lang. Acquis</italic>.</source> <volume>18</volume>, <fpage>433</fpage>&#x02013;<lpage>473</lpage>. <pub-id pub-id-type="doi">10.1017/S0272263100015370</pub-id></citation>
</ref>
<ref id="B27">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hothorn</surname> <given-names>T.</given-names></name> <name><surname>Bretz</surname> <given-names>F.</given-names></name> <name><surname>Westfall</surname> <given-names>P.</given-names></name></person-group> (<year>2008</year>). <article-title>Simultaneous Inference in General Parametric Models</article-title>. <source>Biom. J.</source> <volume>50</volume>, <fpage>346</fpage>&#x02013;<lpage>363</lpage>. <pub-id pub-id-type="doi">10.1002/bimj.200810425</pub-id><pub-id pub-id-type="pmid">18481363</pub-id></citation></ref>
<ref id="B28">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kaschak</surname> <given-names>M. P.</given-names></name> <name><surname>Glenberg</surname> <given-names>A. M.</given-names></name></person-group> (<year>2000</year>). <article-title>Constructing meaning: The role of affordances and grammatical constructions in sentence comprehension</article-title>. <source>J. Mem. Lang.</source> <volume>43</volume>, <fpage>508</fpage>&#x02013;<lpage>529</lpage>. <pub-id pub-id-type="doi">10.1006/jmla.2000.2705</pub-id></citation>
</ref>
<ref id="B29">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kaschak</surname> <given-names>M. P.</given-names></name> <name><surname>Madden</surname> <given-names>C. J.</given-names></name> <name><surname>Therriault</surname> <given-names>D. J.</given-names></name> <name><surname>Yaxley</surname> <given-names>R. H.</given-names></name> <name><surname>Aveyard</surname> <given-names>M.</given-names></name> <name><surname>Blanchard</surname> <given-names>A. A.</given-names></name> <etal/></person-group>. (<year>2005</year>). <article-title>Perception of motion affects language processing</article-title>. <source>Cognition</source> <volume>94</volume>, <fpage>79</fpage>&#x02013;<lpage>89</lpage>. <pub-id pub-id-type="doi">10.1016/j.cognition.2004.06.005</pub-id><pub-id pub-id-type="pmid">15617669</pub-id></citation></ref>
<ref id="B30">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kiefer</surname> <given-names>M.</given-names></name> <name><surname>Pulverm&#x000FC;ller</surname> <given-names>F.</given-names></name></person-group> (<year>2012</year>). <article-title>Conceptual representations in mind and brain: theoretical developments, current evidence and future directions</article-title>. <source>Cortex</source> <volume>48</volume>, <fpage>805</fpage>&#x02013;<lpage>825</lpage>. <pub-id pub-id-type="doi">10.1016/j.cortex.2011.04.006</pub-id><pub-id pub-id-type="pmid">21621764</pub-id></citation></ref>
<ref id="B31">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kronrod</surname> <given-names>A.</given-names></name> <name><surname>Ackerman</surname> <given-names>J. M.</given-names></name></person-group> (<year>2021</year>). <article-title>Under-standing: How embodied states shape inference-making</article-title>. <source>Acta Psychol.</source> <volume>21</volume>:<fpage>103276</fpage>. <pub-id pub-id-type="doi">10.1016/j.actpsy.2021.103276</pub-id><pub-id pub-id-type="pmid">33689912</pub-id></citation></ref>
<ref id="B32">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lakens</surname> <given-names>D.</given-names></name> <name><surname>Caldwell</surname> <given-names>A. R.</given-names></name></person-group> (<year>2021</year>). <article-title>Simulation-based power analysis for factorial analysis of variance designs</article-title>. <source>Adv. Methods Pract. Psychol. Sci.</source> <volume>4</volume>:<fpage>251524592095150</fpage>. <pub-id pub-id-type="doi">10.1177/2515245920951503</pub-id></citation>
</ref>
<ref id="B33">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lindsay</surname> <given-names>S.</given-names></name> <name><surname>Scheepers</surname> <given-names>C.</given-names></name> <name><surname>Kamide</surname> <given-names>Y.</given-names></name></person-group> (<year>2013</year>). <article-title>To dash or to dawdle: verb-associated speed of motion influences eye movements during spoken sentence comprehension</article-title>. <source>PLoS ONE</source> <volume>8</volume>:<fpage>6</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0067187</pub-id><pub-id pub-id-type="pmid">23805299</pub-id></citation></ref>
<ref id="B34">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Magliano</surname> <given-names>J. P.</given-names></name> <name><surname>Zwaan</surname> <given-names>R. A.</given-names></name> <name><surname>Graesser</surname> <given-names>A.</given-names></name></person-group> (<year>1999</year>). <article-title>&#x0201C;The role of situational continuity in narrative understanding,&#x0201D;</article-title> in <source>The Construction of Mental Representations During Reading</source>, eds. H. van Oostendorp and S. R. Goldman (<publisher-loc>New Jersey</publisher-loc>: <publisher-name>Lawrence Erlbaum Associates Publishers</publisher-name>), <fpage>219</fpage>&#x02013;<lpage>245</lpage>.</citation>
</ref>
<ref id="B35">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Moreno</surname> <given-names>I.</given-names></name> <name><surname>De Vega</surname> <given-names>M.</given-names></name> <name><surname>Le&#x000F3;n</surname> <given-names>I.</given-names></name></person-group> (<year>2013</year>). <article-title>Understanding action language modulates oscillatory mu and beta rhythms in the same way as observing actions</article-title>. <source>Brain Cogn.</source> <volume>82</volume>, <fpage>236</fpage>&#x02013;<lpage>242</lpage>. <pub-id pub-id-type="doi">10.1016/j.bandc.2013.04.010</pub-id><pub-id pub-id-type="pmid">23711935</pub-id></citation></ref>
<ref id="B36">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Norman</surname> <given-names>T.</given-names></name> <name><surname>Peleg</surname> <given-names>O.</given-names></name></person-group> (<year>2022</year>). <article-title>The reduced embodiment of a second language</article-title>. <source>Bilingualism: Lang. Cogn.</source> <volume>25</volume>, <fpage>406</fpage>&#x02013;<lpage>416</lpage>. <pub-id pub-id-type="doi">10.1017/S1366728921001115</pub-id></citation>
</ref>
<ref id="B37">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Pavlenko</surname> <given-names>A.</given-names></name></person-group> (<year>2014</year>). <source>The Bilingual Mind: And What it Tells us About Language and Thought</source>. <publisher-loc>New York</publisher-loc>: <publisher-name>Cambridge University Press</publisher-name>.</citation>
</ref>
<ref id="B38">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>P&#x000E9;rez</surname> <given-names>A.</given-names></name> <name><surname>Hansen</surname> <given-names>L.</given-names></name> <name><surname>Bajo</surname> <given-names>T.</given-names></name></person-group> (<year>2019</year>). <article-title>The nature of first and second language processing: The role of cognitive control and L2 proficiency during text-level comprehension</article-title>. <source>Biling.</source> <volume>22</volume>, <fpage>930</fpage>&#x02013;<lpage>948</lpage>. <pub-id pub-id-type="doi">10.1017/S1366728918000846</pub-id></citation>
</ref>
<ref id="B39">
<citation citation-type="web"><person-group person-group-type="author"><name><surname>Pinheiro</surname> <given-names>J.</given-names></name> <name><surname>Bates</surname> <given-names>D.</given-names></name> <name><surname>DebRoy</surname> <given-names>S.</given-names></name> <name><surname>Sarkar</surname> <given-names>D.</given-names></name> <collab>R Core Team</collab></person-group> (<year>2021</year>). <source>nlme: Linear and Nonlinear Mixed Effects Models</source>. Available at: <ext-link ext-link-type="uri" xlink:href="https://CRAN.R-project.org/package=nlme">https://CRAN.R-project.org/package=nlme</ext-link> (accessed November 30, 2019).</citation>
</ref>
<ref id="B40">
<citation citation-type="web"><person-group person-group-type="author"><collab>R Core Team</collab></person-group> (<year>2019</year>). <source>R: A Language and Environment for Statistical Computing.</source> Available at: <ext-link ext-link-type="uri" xlink:href="http://www.R-project.org/">http://www.R-project.org/</ext-link> (accessed November 30, 2019).</citation>
</ref>
<ref id="B41">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Santana</surname> <given-names>E. J.</given-names></name> <name><surname>De Vega</surname> <given-names>M.</given-names></name></person-group> (<year>2013</year>). <article-title>An ERP study of motor compatibility effects in action language</article-title>. <source>Brain Res.</source> <volume>1526</volume>, <fpage>71</fpage>&#x02013;<lpage>83</lpage>. <pub-id pub-id-type="doi">10.1016/j.brainres.2013.06.020</pub-id><pub-id pub-id-type="pmid">23796780</pub-id></citation></ref>
<ref id="B42">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schuil</surname> <given-names>K. D.</given-names></name> <name><surname>Smits</surname> <given-names>M.</given-names></name> <name><surname>Zwaan</surname> <given-names>R. A.</given-names></name></person-group> (<year>2013</year>). <article-title>Sentential context modulates the involvement of the motor cortex in action language processing: an fMRI study</article-title>. <source>Front. Hum Neurosci</source>. <volume>7</volume>:<fpage>100</fpage>. <pub-id pub-id-type="doi">10.3389/fnhum.2013.00100</pub-id><pub-id pub-id-type="pmid">23580364</pub-id></citation></ref>
<ref id="B43">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Shapiro</surname> <given-names>L.</given-names></name></person-group> (<year>2019</year>). <source>Embodied Cognition, 2nd Edn</source>. <publisher-loc>London</publisher-loc>: <publisher-name>Routledge</publisher-name>.</citation>
</ref>
<ref id="B44">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shiang</surname> <given-names>R. F.</given-names></name></person-group> (<year>2018</year>). <article-title>Embodied EFL reading activity: let&#x00027;s produce comics</article-title>. <source>Read. Foreign Lang.</source> <volume>30</volume>, <fpage>108</fpage>&#x02013;<lpage>129</lpage>.</citation>
</ref>
<ref id="B45">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Speed</surname> <given-names>L. J.</given-names></name> <name><surname>Vigliocco</surname> <given-names>G.</given-names></name></person-group> (<year>2014</year>). <article-title>Eye movements reveal the dynamic simulation of speed in language</article-title>. <source>Cogn. Sci.</source> <volume>38</volume>, <fpage>367</fpage>&#x02013;<lpage>382</lpage>. <pub-id pub-id-type="doi">10.1111/cogs.12096</pub-id><pub-id pub-id-type="pmid">24795958</pub-id></citation></ref>
<ref id="B46">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Spivey</surname> <given-names>M.</given-names></name> <name><surname>Richardson</surname> <given-names>D.</given-names></name></person-group> (<year>2009</year>). <article-title>&#x0201C;Language processing embodied and embedded,&#x0201D;</article-title> in <source>The Cambridge Handbook of Situated Cognition</source>, eds. P. Robbins and M. Aydede (<publisher-loc>New York</publisher-loc>: <publisher-name>Cambridge University Press</publisher-name>), <fpage>382</fpage>&#x02013;<lpage>400</lpage>.</citation>
</ref>
<ref id="B47">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tettamanti</surname> <given-names>M.</given-names></name> <name><surname>Buccino</surname> <given-names>G.</given-names></name> <name><surname>Saccuman</surname> <given-names>M. C.</given-names></name> <name><surname>Gallese</surname> <given-names>V.</given-names></name> <name><surname>Danna</surname> <given-names>M.</given-names></name> <name><surname>Scifo</surname> <given-names>P.</given-names></name> <etal/></person-group>. (<year>2005</year>). <article-title>Listening to action-related sentences activates fronto-parietal motor circuits</article-title>. <source>J. Cogn. Neurosci.</source> <volume>17</volume>, <fpage>273</fpage>&#x02013;<lpage>281</lpage>. <pub-id pub-id-type="doi">10.1162/0898929053124965</pub-id><pub-id pub-id-type="pmid">15811239</pub-id></citation></ref>
<ref id="B48">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tian</surname> <given-names>L.</given-names></name> <name><surname>Chen</surname> <given-names>H.</given-names></name> <name><surname>Zhao</surname> <given-names>W.</given-names></name> <name><surname>Wu</surname> <given-names>J.</given-names></name> <name><surname>Zhang</surname> <given-names>Q.</given-names></name> <name><surname>De</surname> <given-names>A.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>The role of motor system in action-related language comprehension in L1 and L2: an fMRI study</article-title>. <source>Brain Lang</source>. <volume>201</volume>:<fpage>104714</fpage>. <pub-id pub-id-type="doi">10.1016/j.bandl.2019.104714</pub-id><pub-id pub-id-type="pmid">31790907</pub-id></citation></ref>
<ref id="B49">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ulbricht</surname> <given-names>J. N.</given-names></name></person-group> (<year>2020</year>). <article-title>The embodied teaching of spatial terms: gestures mapped to morphemes improve learning</article-title>. <source>Front. Educ.</source> 5<italic>:</italic>109. <pub-id pub-id-type="doi">10.3389/feduc.2020.00109</pub-id></citation>
</ref>
<ref id="B50">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Vukovic</surname> <given-names>N.</given-names></name> <name><surname>Shtyrov</surname> <given-names>Y.</given-names></name></person-group> (<year>2014</year>). <article-title>Cortical motor systems are involved in second language comprehension: evidence from rapid mu-rhythm desynchronisation</article-title>. <source>Neuroimage</source> <volume>102</volume>, <fpage>695</fpage>&#x02013;<lpage>703</lpage>. <pub-id pub-id-type="doi">10.1016/j.neuroimage.2014.08.039</pub-id><pub-id pub-id-type="pmid">25175538</pub-id></citation></ref>
<ref id="B51">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Vukovic</surname> <given-names>N.</given-names></name> <name><surname>Williams</surname> <given-names>J. N.</given-names></name></person-group> (<year>2014</year>). <article-title>Automatic perceptual simulation of first language meanings during second language sentence processing in bilinguals</article-title>. <source>Acta Psychol.</source> <volume>145</volume>, <fpage>98</fpage>&#x02013;<lpage>103</lpage>. <pub-id pub-id-type="doi">10.1016/j.actpsy.2013.11.002</pub-id><pub-id pub-id-type="pmid">24333464</pub-id></citation></ref>
<ref id="B52">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zhang</surname> <given-names>H.</given-names></name> <name><surname>Vanek</surname> <given-names>N.</given-names></name></person-group> (<year>2021</year>). <article-title>From &#x02018;No, she does&#x00027; to &#x02018;Yes, she does&#x00027;: negation processing in negative yes-no questions by Mandaring speakers of English</article-title>. <source>Appl Psycholinguist.</source> <volume>42</volume>, <fpage>937</fpage>&#x02013;<lpage>967</lpage>. <pub-id pub-id-type="doi">10.1017/S0142716421000175</pub-id></citation>
</ref>
<ref id="B53">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zwaan</surname> <given-names>A.</given-names></name></person-group> (<year>1999</year>). <article-title>Embodied cognition, perceptual symbols, and situation models</article-title>. <source>Discour. Proc</source>. <volume>28</volume>, <fpage>81</fpage>&#x02013;<lpage>88</lpage>. <pub-id pub-id-type="doi">10.1080/01638539909545070</pub-id></citation>
</ref>
<ref id="B54">
<citation citation-type="book"><person-group person-group-type="author"><name><surname>Zwaan</surname> <given-names>R. A.</given-names></name></person-group> (<year>2004</year>). <article-title>&#x0201C;The immersed experiencer: toward an embodied theory of language comprehension,&#x0201D;</article-title> in <source>The psychology of learning and motivation: Advances in research and theory</source>, Eds. B. H. Ross (<publisher-loc>London</publisher-loc>: <publisher-name>Elsevier Science</publisher-name>), <fpage>35</fpage>&#x02013;<lpage>62</lpage>.</citation>
</ref>
<ref id="B55">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zwaan</surname> <given-names>R. A.</given-names></name></person-group> (<year>2014</year>). <article-title>Embodiment and language comprehension: reframing the discussion</article-title>. <source>Trends Cogn. Sci</source>. <volume>18</volume>, <fpage>229</fpage>&#x02013;<lpage>234</lpage>. <pub-id pub-id-type="doi">10.1016/j.tics.2014.02.008</pub-id><pub-id pub-id-type="pmid">24630873</pub-id></citation></ref>
<ref id="B56">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zwaan</surname> <given-names>R. A.</given-names></name> <name><surname>Brown</surname> <given-names>C. M.</given-names></name></person-group> (<year>1996</year>). <article-title>The influence of language proficiency and comprehension skill on situation-model construction</article-title>. <source>Discour. Proc.</source> <volume>21</volume>, <fpage>289</fpage>&#x02013;<lpage>327</lpage>. <pub-id pub-id-type="doi">10.1080/01638539609544960</pub-id></citation>
</ref>
<ref id="B57">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zwaan</surname> <given-names>R. A.</given-names></name> <name><surname>Pecher</surname> <given-names>D.</given-names></name></person-group> (<year>2012</year>). <article-title>Revisiting mental simulation in language comprehension: Six replication attempts</article-title>. <source>PLoS ONE</source> <volume>7</volume>:<fpage>e0051382</fpage>. <pub-id pub-id-type="doi">10.1371/journal.pone.0051382</pub-id><pub-id pub-id-type="pmid">23300547</pub-id></citation></ref>
<ref id="B58">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zwaan</surname> <given-names>R. A.</given-names></name> <name><surname>Radvansky</surname> <given-names>G. A.</given-names></name></person-group> (<year>1998</year>). <article-title>Situation models in language comprehension and memory</article-title>. <source>Psychol. Bull.</source> <volume>123</volume>:<fpage>162</fpage>. <pub-id pub-id-type="doi">10.1037/0033-2909.123.2.162</pub-id><pub-id pub-id-type="pmid">9522683</pub-id></citation></ref>
<ref id="B59">
<citation citation-type="journal"><person-group person-group-type="author"><name><surname>Zwaan</surname> <given-names>R. A.</given-names></name> <name><surname>Taylor</surname> <given-names>L. J.</given-names></name></person-group> (<year>2006</year>). <article-title>Seeing, acting, understanding: motor resonance in language comprehension</article-title>. <source>J. Exp. Psychol. Gen.</source> <volume>135</volume>, <fpage>1</fpage>&#x02013;<lpage>11</lpage>. <pub-id pub-id-type="doi">10.1037/0096-3445.135.1.1</pub-id><pub-id pub-id-type="pmid">16478313</pub-id></citation></ref>
</ref-list>
</back>
</article>