<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Psychol.</journal-id>
<journal-title>Frontiers in Psychology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Psychol.</abbrev-journal-title>
<issn pub-type="epub">1664-1078</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpsyg.2025.1665380</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Psychology</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Material hardship, not household income, predicts impaired punishment learning: a computational reinforcement learning perspective</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Wang</surname>
<given-names>Zhen</given-names>
</name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>He</surname>
<given-names>Xu</given-names>
</name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/3133200/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Su</surname>
<given-names>Yunsheng</given-names>
</name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Bu</surname>
<given-names>Laijun</given-names>
</name>
<xref ref-type="aff" rid="aff5"><sup>5</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Wang</surname>
<given-names>Yi</given-names>
</name>
<xref ref-type="aff" rid="aff6"><sup>6</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Guangzhou Xinhua University</institution>, <addr-line>Dongguan</addr-line>, <country>China</country></aff>
<aff id="aff2"><sup>2</sup><institution>School of Public Health and Management, Guangzhou University of Chinese Medicine</institution>, <addr-line>Guangzhou</addr-line>, <country>China</country></aff>
<aff id="aff3"><sup>3</sup><institution>School of Psychology, South China Normal Univeristy</institution>, <addr-line>Guangzhou</addr-line>, <country>China</country></aff>
<aff id="aff4"><sup>4</sup><institution>School of Journalism and Communication, Jinan University</institution>, <addr-line>Guangzhou</addr-line>, <country>China</country></aff>
<aff id="aff5"><sup>5</sup><institution>School of Nursing, Guangdong Pharmaceutical University</institution>, <addr-line>Guangzhou</addr-line>, <country>China</country></aff>
<aff id="aff6"><sup>6</sup><institution>School of Journalism and Communication, Guangzhou University</institution>, <addr-line>Guangzhou</addr-line>, <country>China</country></aff>
<author-notes>
<fn fn-type="edited-by" id="fn0001">
<p>Edited by: <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1993119/overview">Yanfeng Xu</ext-link>, University of South Carolina, United States</p>
</fn>
<fn fn-type="edited-by" id="fn0002">
<p>Reviewed by: <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/139582/overview">Jeffrey Coldren</ext-link>, Youngstown State University, United States</p>
<p><ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/3193492/overview">Huaiyu Liu</ext-link>, University College London, United Kingdom</p>
</fn>
<corresp id="c001">&#x002A;Correspondence: Yi Wang, <email>309394931@qq.com</email></corresp>
</author-notes>
<pub-date pub-type="epub">
<day>22</day>
<month>10</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>16</volume>
<elocation-id>1665380</elocation-id>
<history>
<date date-type="received">
<day>14</day>
<month>07</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>01</day>
<month>10</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2025 Wang, He, Su, Bu and Wang.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Wang, He, Su, Bu and Wang</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<sec>
<title>Introduction</title>
<p>Socioeconomic disadvantage has been linked to neurocognitive alterations in reward and loss processing, which may contribute to adverse psychological outcomes. However, the mechanisms through which it influences reinforcement learning remain unclear.</p>
</sec>
<sec>
<title>Methods</title>
<p>This study employed a Probabilistic Reversal Learning Task to examine how two distinct indicators of disadvantage&#x2014;material hardship and low household income&#x2014;affect reward and punishment-based learning in a sample of Chinese undergraduate students. Behavioral responses were analyzed through computational modeling within a reinforcement learning framework, estimating three key parameters: reward learning rate, punishment learning rate, and inverse temperature.</p>
</sec>
<sec>
<title>Results</title>
<p>Results revealed that material hardship uniquely predicted individual differences in punishment learning rate, whereas household income showed no independent association with any of the model parameters.</p>
</sec>
<sec>
<title>Discussion</title>
<p>The findings suggest that material hardship may specifically impair the ability to learn from negative outcomes. Furthermore, the study underscores the importance of distinguishing between material hardship and income-based adversity in research examining the cognitive impacts of socioeconomic disadvantage.</p>
</sec>
</abstract>
<kwd-group>
<kwd>material hardship</kwd>
<kwd>socioeconomic disadvantage</kwd>
<kwd>reinforcement learning</kwd>
<kwd>punishment learning</kwd>
<kwd>computational modeling</kwd>
</kwd-group>
<counts>
<fig-count count="3"/>
<table-count count="2"/>
<equation-count count="3"/>
<ref-count count="40"/>
<page-count count="8"/>
<word-count count="5894"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Human Developmental Psychology</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec1">
<label>1</label>
<title>Introduction</title>
<p>In the field of cognitive science, reinforcement learning (RL) refers to a fundamental cognitive process by which individuals optimize their behavior based on environmental feedback (<xref ref-type="bibr" rid="ref30">Shteingart and Loewenstein, 2014</xref>; <xref ref-type="bibr" rid="ref33">Subramanian et al., 2022</xref>). This process operates via two dissociable pathways: (1) reward learning, where actions may be strengthened by positive outcomes (<xref ref-type="bibr" rid="ref6">Daniel and Pollmann, 2014</xref>)&#x2014;for example, an employee works harder after receiving a bonus&#x2014;and (2) punishment learning, where behaviors may be modified to avoid adverse consequences, such as a driver slowing down after receiving a speeding ticket. Neuroscience research indicates that these pathways engage distinct neural substrates (<xref ref-type="bibr" rid="ref40">Yacubian et al., 2006</xref>; <xref ref-type="bibr" rid="ref39">Xue et al., 2013</xref>). Critically, extensive research has demonstrated that reward and punishment learning plays a crucial role in everyday decision-making (<xref ref-type="bibr" rid="ref22">Lee et al., 2012</xref>), influencing behaviors across diverse contexts ranging from risk-taking (<xref ref-type="bibr" rid="ref26">Marshall and Kirkpatrick, 2017</xref>) to social interactions (<xref ref-type="bibr" rid="ref18">Heininga et al., 2017</xref>). This framework helps explain socioeconomic disparities in behavior; for instance, higher socioeconomic status has been linked to risky driving behaviors (<xref ref-type="bibr" rid="ref2">Atombo et al., 2017</xref>), potentially because the punitive impact of fines is attenuated, disrupting the typical balance of punishment learning. While the behavioral and neural mechanisms of RL are well-documented, few studies investigate how individual differences, such as early-life experiences, influence these mechanisms. Investigating such factors may clarify the determinants of lifelong learning tendencies, thereby integrating cognitive models of decision-making with developmental psychology.</p>
<p>Given the established role of RL in daily life, a critical yet understudied question is how socioeconomic factors&#x2014;particularly socioeconomic disadvantage&#x2014;may shape these cognitive processes. Socioeconomic disadvantage exerts profound and far-reaching influences on human development, with measurable effects across multiple life domains including physical health (<xref ref-type="bibr" rid="ref35">Torpy et al., 2007</xref>), mental well-being (<xref ref-type="bibr" rid="ref25">Marbin et al., 2022</xref>), cognitive functioning (<xref ref-type="bibr" rid="ref24">Mani et al., 2013</xref>), and economic decision-making (<xref ref-type="bibr" rid="ref9">De Bruijn and Antonides, 2022</xref>). Notably, emerging neuroimaging evidence indicates that socioeconomic disadvantage may alter neurocognitive mechanisms relevant to RL, such as reward and loss processing. For example, <xref ref-type="bibr" rid="ref38">White et al. (2022)</xref> found that a lower income-to-poverty ratio was associated with heightened neural responses to reward and loss cues during a passive avoidance task. <xref ref-type="bibr" rid="ref28">Romens et al. (2015)</xref> demonstrated that increased neural activity during reward anticipation mediated the association between childhood poverty and depression symptoms, suggesting a potential neural pathway linking socioeconomic disadvantage to mental health outcomes. However, despite these advances, direct evidence on whether and how socioeconomic disadvantage modulates RL processes remains scarce. Addressing this gap could not only bridge cognitive science with developmental psychology but also inform interventions to mitigate the long-term behavioral impacts of socioeconomic disadvantage.</p>
<p>Over the past two decades, researchers have increasingly examined material hardship as a proximal measure of socioeconomic disadvantage (<xref ref-type="bibr" rid="ref14">Gershoff et al., 2007</xref>; <xref ref-type="bibr" rid="ref34">Thomas and Waldfogel, 2022</xref>). Unlike conventional income-based measures, material hardship reflects tangible deficits in meeting basic needs&#x2014;such as food insecurity, unstable housing, and lack of medical care&#x2014;providing a proximate framework to examine how acute scarcity shapes cognition and behavior (<xref ref-type="bibr" rid="ref4">Beverly, 2001</xref>). Recent studies suggest that these experiences may influence economic decision-making, potentially altering how individuals evaluate risks and rewards. For example, <xref ref-type="bibr" rid="ref17">He et al. (2024)</xref> reported that individuals with higher material hardship exhibited more loss-averse behavior in a mixed gambling task. Additionally, neuroimaging evidence demonstrates associations between material hardship and functional changes in frontal-limbic circuit (<xref ref-type="bibr" rid="ref5">Chen et al., 2023</xref>), which is also a neural network critically involved in RL processes. These observations raise the possibility that material hardship, as a concrete manifestation of socioeconomic disadvantage, may directly modulate RL mechanisms, exacerbating maladaptive decision-making. By integrating material hardship into cognitive psychology, we can bridge the gap between macro-level socioeconomic factors and micro-level cognitive processes, ultimately clarifying how specific deprivation experiences shapes long-term behavior.</p>
<p>To empirically examine RL processes, researchers often employ probabilistic learning tasks (<xref ref-type="bibr" rid="ref21">Koch et al., 2008</xref>; <xref ref-type="bibr" rid="ref7">Daniel et al., 2020</xref>). In these paradigms, participants learn through trial and error to associate actions with probabilistically delivered rewards or punishments, thereby capturing adaptive learning under uncertainty (<xref ref-type="bibr" rid="ref31">Soltani and Izquierdo, 2019</xref>). Computational RL models are then used to quantify the latent learning processes and individual differences (<xref ref-type="bibr" rid="ref29">Schaaf et al., 2023</xref>). These models mathematically describe how individuals update their expectations based on feedback received, enabling the estimation of parameters reflecting distinct cognitive components. Key parameters include the learning rate, which determines how quickly expectations adjust to new feedback, and inverse temperature, which indicates the degree of randomness in decision-making (<xref ref-type="bibr" rid="ref20">Katahira, 2015</xref>). Critically, while standard RL models apply a single learning rate to both reward and punishment outcomes, evidence from cognitive neuroscience research suggests dissociable neural substrates for these processes (<xref ref-type="bibr" rid="ref16">Gueguen et al., 2021</xref>). This supports the use of a three-parameter model decoupling reward and punishment learning (<xref ref-type="bibr" rid="ref10">den Ouden et al., 2013</xref>): the reward learning rate determines how rapidly expectations increase following gains, the punishment learning rate governs how rapidly expectations decrease following losses, and the inverse temperature parameter captures choice stochasticity.</p>
<p>Building upon this foundation and addressing the identified research gap, the current study employs a probabilistic reversal learning task coupled with the three-parameter computational RL model to empirically test whether socioeconomic disadvantage modulates core RL mechanisms. Specifically, we examine how two established indicators of disadvantage&#x2014;material hardship and low household income&#x2014;influence the efficiency of learning. These indicators are included as independent variables in regression analyses to assess their effects on two key computational parameters: the reward learning rate and the punishment learning rate. Based on emerging neurocognitive evidence linking socioeconomic adversity to heightened neural sensitivity to rewards and punishments (<xref ref-type="bibr" rid="ref38">White et al., 2022</xref>), we hypothesized that greater socioeconomic disadvantage will be associated with elevated learning rates for both rewarding and punishing outcomes. This accelerated behavioral adaptation to feedback represents a potential cognitive mechanism through which socioeconomic disadvantage could shape long-term decision-making tendencies. By employing computational modeling within this well-established RL paradigm, our study moves beyond behavioral correlations to directly probe how disadvantage modulates these learning mechanisms, thereby illuminating cognitive pathways linking socioeconomic context to adaptive decision-making.</p>
</sec>
<sec sec-type="materials|methods" id="sec2">
<label>2</label>
<title>Materials and methods</title>
<sec id="sec3">
<label>2.1</label>
<title>Participants</title>
<p>The study protocol received ethical approval from the Research Ethics Committee of the author&#x2019;s affiliated university. <italic>A priori</italic> power analysis using G&#x002A;Power (<xref ref-type="bibr" rid="ref11">Faul et al., 2007</xref>) indicated that a sample size of 84 provided 95% power to detect small effects (0.2) in multiple regression with up to 4 predictors at <italic>&#x03B1;</italic>&#x202F;=&#x202F;0.05. A total of 100 first-year undergraduates were recruited from a public comprehensive university in China, where the average scores on the National College Entrance Examination (Gaokao) of admitted students fall within the mid-to-upper range nationally. Following exclusions for incomplete data or task accuracy below chance level, 95 participants (57 females, 38 males; aged 18&#x2013;20&#x202F;years, <italic>M</italic>&#x202F;&#x00B1;&#x202F;<italic>SD</italic>&#x202F;=&#x202F;18.44&#x202F;&#x00B1;&#x202F;0.58) comprised the final sample. All participants reported normal or corrected-to-normal vision, and none reported a history of psychotropic medication use. Written informed consent was obtained prior to participation. Participants received &#x00A5;30&#x2013;50 (approximately 5&#x2013;6 USD) as compensation for their time.</p>
</sec>
<sec id="sec4">
<label>2.2</label>
<title>Measures</title>
<p><italic>Material hardship</italic>. Material hardship was assessed using the Chinese version of the Family Economic Hardship Questionnaire (<xref ref-type="bibr" rid="ref37">Wang et al., 2010</xref>). The 4-item scale evaluates the frequency of material hardships across four domains: food insecurity, clothing affordability, access to entertainment, and housing stability. It has demonstrated strong psychometric properties in Chinese adolescent samples, with a Cronbach&#x2019;s <italic>&#x03B1;</italic> of 0.84 in the original study and 0.83 in our sample. Participants rated each item on a 5-point Likert scale (1&#x202F;=&#x202F;never, 5&#x202F;=&#x202F;all the time). A composite score was calculated by averaging responses, with higher scores indicating greater material hardship.</p>
<p><italic>Household income</italic>. Household income was self-reported using a 7-point ordinal scale: 1 (monthly income &#x003C; &#x00A5;4,000 [&#x2248;5,060 USD]), 2 (&#x00A5;4,000&#x2013;7,999 [&#x2248;560&#x2013;1,100 USD]), 3 (&#x00A5;8,000&#x2013;11,999 [&#x2248;1,100&#x2013;1,680 USD]), 4 (&#x00A5;12,000&#x2013;15,999 [&#x2248;1,680&#x2013;2,230 USD]), 5 (&#x00A5;16,000&#x2013;19,999 [&#x2248;2,230&#x2013;2,800 USD]), 6 (&#x00A5;20,000&#x2013;39,999 [&#x2248;2,800&#x2013;5,600 USD]), to 7 (&#x2265;&#x00A5;40,000 [&#x2248;5,600 USD]), with lower scores indicating lower household income.</p>
<p><italic>Probabilistic reversal learning task</italic>. Participants performed a computerized probabilistic reversal learning task (adapted from <xref ref-type="bibr" rid="ref15">Gl&#x00E4;scher et al., 2009</xref>) designed to measure reinforcement learning mechanisms under uncertainty. In this task, participants repeatedly selected between two visual stimuli&#x2014;a square and a circle&#x2014;presented simultaneously on each trial, with the goal of maximizing monetary rewards. They were explicitly informed that accumulated winnings would supplement their base compensation. Each trial followed a structured sequence: Following stimulus onset, participants had 1,500&#x202F;ms to select one option; failure to respond within this window triggered an automatic random selection by the computer, with reaction time recorded as 1,500&#x202F;ms. The chosen stimulus was then highlighted for 500&#x202F;ms. After a variable delay (500&#x2013;1,500&#x202F;ms), the outcome (WIN &#x00A5;0.5 or LOSS &#x00A5;0) was displayed for 1,000&#x202F;ms. Critically, stimulus-outcome contingencies were probabilistic: One stimulus was designated &#x201C;correct&#x201D; (75% probability of WIN; 25% probability of LOSS), while the other was &#x201C;incorrect&#x201D; (25% WIN; 75% LOSS). The &#x201C;LOSS&#x201D; outcome was coded as &#x00A5;0.00 (instead of a negative value) to avoid negative earnings throughout the task. This design was implemented to maintain participant motivation and engagement, and although the outcome is numerically zero, it is psychologically perceived as a loss relative to the winning outcome. A variable inter-trial interval (500&#x2013;1,500&#x202F;ms) followed, resulting in a mean trial duration of approximately 5,000&#x202F;ms (see <xref ref-type="fig" rid="fig1">Figure 1</xref> for schematic). To assess adaptive learning, the contingencies reversed randomly after 5 or 6 consecutive correct choices (this variability prevented anticipation of reversals). Participants needed to learn the new contingencies before another reversal could occur. Across the 60-trial task, up to 10 reversals were possible, with the total number of achieved reversals serving as a behavioral index of adaptive learning capacity. The computational model was fitted to each participant&#x2019;s trial-by-trial choice data (i.e., which stimulus was selected on each trial), along with the corresponding outcomes (win or loss). Aggregate measures such as accuracy and reversal frequency were used solely as behavioral indices of task performance.</p>
<fig position="float" id="fig1">
<label>Figure 1</label>
<caption>
<p>The procedure of a single trial in the probabilistic reversal learning task.</p>
</caption>
<graphic xlink:href="fpsyg-16-1665380-g001.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Flowchart of a single trial in a probabilistic reversal learning task, showing the sequence and timings: Response Window (1,500 ms), Choice (500 ms), Delay (500&#x2013;1,500 ms), Outcome (1,000 ms), and Inter-trial Interval (500&#x2013;1,500 ms). Stimuli include a square and a circle, with outcomes shown as a coin (win) or a gray circle (loss).</alt-text>
</graphic>
</fig>
</sec>
<sec id="sec5">
<label>2.3</label>
<title>Computational modeling of reinforcement learning</title>
<p>While several computational frameworks exist for modeling reinforcement learning, we selected the three-parameter model (<xref ref-type="bibr" rid="ref10">den Ouden et al., 2013</xref>) for its theoretical alignment with our research questions. This model distinguishes between reward and punishment learning rates, capturing dissociable mechanisms in belief updating. The model operates on a trial-by-trial basis. First, the prediction error (<inline-formula>
<mml:math id="M1">
<mml:msub>
<mml:mi>PE</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
</mml:math>
</inline-formula>) is calculated as shown in <xref ref-type="disp-formula" rid="EQ1">Equation (1)</xref>:</p>
<disp-formula id="EQ1">
<label>(1)</label>
<mml:math id="M2">
<mml:msub>
<mml:mi>PE</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>EV</mml:mi>
<mml:mrow>
<mml:mi mathvariant="normal">t</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:math>
</disp-formula>
<p>where <inline-formula>
<mml:math id="M3">
<mml:msub>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> is the outcome (scaled to 1 for win, &#x2212;1 for loss) and <inline-formula>
<mml:math id="M4">
<mml:msub>
<mml:mi>EV</mml:mi>
<mml:mrow>
<mml:mi mathvariant="normal">t</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:math>
</inline-formula> is the expected value from the previous trial (initialized to 0 at <italic>t</italic>&#x202F;=&#x202F;1). Then, the expected value for the chosen stimulus at trial <italic>t</italic> (<inline-formula>
<mml:math id="M5">
<mml:msub>
<mml:mi>EV</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
</mml:math>
</inline-formula>) is updated using this prediction error according to <xref ref-type="disp-formula" rid="EQ2">Equation (2)</xref>:</p>
<disp-formula id="EQ2">
<label>(2)</label>
<mml:math id="M6">
<mml:msub>
<mml:mi>EV</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:msub>
<mml:mi>EV</mml:mi>
<mml:mrow>
<mml:mi mathvariant="normal">t</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>+</mml:mo>
<mml:mi>&#x03B1;</mml:mi>
<mml:mo>&#x00D7;</mml:mo>
<mml:msub>
<mml:mi>PE</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
</mml:math>
</disp-formula>
<p>This update is governed by separate learning rates (<inline-formula>
<mml:math id="M7">
<mml:mi>&#x03B1;</mml:mi>
</mml:math>
</inline-formula>) for positive and negative prediction errors; specifically, specifically, the reward learning rate (<inline-formula>
<mml:math id="M8">
<mml:msup>
<mml:mi>&#x03B1;</mml:mi>
<mml:mo>+</mml:mo>
</mml:msup>
</mml:math>
</inline-formula>) is applied when the <inline-formula>
<mml:math id="M9">
<mml:msub>
<mml:mi>PE</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
<mml:mo>&#x003E;</mml:mo>
<mml:mn>0</mml:mn>
</mml:math>
</inline-formula>, while the punishment learning rate (<inline-formula>
<mml:math id="M10">
<mml:msup>
<mml:mi>&#x03B1;</mml:mi>
<mml:mo>&#x2212;</mml:mo>
</mml:msup>
</mml:math>
</inline-formula>) is applied when <inline-formula>
<mml:math id="M11">
<mml:msub>
<mml:mi>PE</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
<mml:mo>&#x2264;</mml:mo>
<mml:mn>0</mml:mn>
</mml:math>
</inline-formula>. This follows the approach in the hBayesDM package (<xref ref-type="bibr" rid="ref1">Ahn et al., 2017</xref>) for this class of models, where a prediction error&#x202F;&#x2264;&#x202F;0 (outcome is worse than or equal to expectation) engages the punishment learning system for updating. Subsequently, the probability (P) of choosing options A and B is determined by a softmax function defined in <xref ref-type="disp-formula" rid="EQ3">Equation (3)</xref>:</p>
<disp-formula id="EQ3">
<label>(3)</label>
<mml:math id="M12">
<mml:mi mathvariant="normal">P</mml:mi>
<mml:mo stretchy="true">(</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">A</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
<mml:mo stretchy="true">)</mml:mo>
<mml:mo>=</mml:mo>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>+</mml:mo>
<mml:msup>
<mml:mi mathvariant="normal">e</mml:mi>
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>&#x03B2;</mml:mi>
<mml:mo>&#x00D7;</mml:mo>
<mml:mo stretchy="true">(</mml:mo>
<mml:msub>
<mml:mi>EV</mml:mi>
<mml:mi mathvariant="normal">A</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>EV</mml:mi>
<mml:mi mathvariant="normal">B</mml:mi>
</mml:msub>
<mml:mo stretchy="true">)</mml:mo>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
<mml:mi mathvariant="normal">P</mml:mi>
<mml:mo stretchy="true">(</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">B</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
<mml:mo stretchy="true">)</mml:mo>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mi mathvariant="normal">P</mml:mi>
<mml:mo stretchy="true">(</mml:mo>
<mml:msub>
<mml:mi mathvariant="normal">A</mml:mi>
<mml:mi mathvariant="normal">t</mml:mi>
</mml:msub>
<mml:mo stretchy="true">)</mml:mo>
</mml:math>
</disp-formula>
<p>where the inverse temperature parameter (&#x03B2;) governs the stochasticity of choices, with higher values indicating more deterministic, value-driven decision-making. Parameters (<inline-formula>
<mml:math id="M13">
<mml:msup>
<mml:mi>&#x03B1;</mml:mi>
<mml:mo>+</mml:mo>
</mml:msup>
</mml:math>
</inline-formula>, <inline-formula>
<mml:math id="M14">
<mml:msup>
<mml:mi>&#x03B1;</mml:mi>
<mml:mo>&#x2212;</mml:mo>
</mml:msup>
</mml:math>
</inline-formula>, <inline-formula>
<mml:math id="M15">
<mml:mi>&#x03B2;</mml:mi>
</mml:math>
</inline-formula>) were estimated for each participant using a hierarchical Bayesian approach implemented in the hBayesDM package (<xref ref-type="bibr" rid="ref1">Ahn et al., 2017</xref>) in R. This method was chosen because it provides more robust estimates by simultaneously modeling individual and group-level parameters, using the group distribution to constrain improbable individual estimates through partial pooling. Model parameters were estimated using Markov Chain Monte Carlo (MCMC) sampling, and convergence was successfully confirmed by R-hat values&#x202F;&#x003C;&#x202F;1.01.</p>
</sec>
<sec id="sec6">
<label>2.4</label>
<title>Statistical analysis</title>
<p>Linear regression models examined how material hardship and household income independently predicted reward learning rate, punishment learning rate, and inverse temperature. Age and gender were included as covariates. Significance was evaluated at <italic>p</italic>&#x202F;&#x003C;&#x202F;0.05, with effect sizes reported as standardized coefficients (<italic>b</italic>). To ensure robustness, we also applied False Discovery Rate (FDR) correction for multiple comparisons across the three primary dependent variables (<xref ref-type="bibr" rid="ref3">Benjamini and Hochberg, 1995</xref>). However, in interpreting the results, we focus on the pattern of effect sizes and their confidence intervals, as these provide more meaningful information than dichotomous significance testing alone.</p>
<p>We estimated three separate linear regression models. In each model, one of the computational parameters (reward learning rate, punishment learning rate, or inverse temperature) served as the dependent variable. The key independent variables of interest&#x2014;material hardship and household income&#x2014;were entered simultaneously into each model, along with the covariates of age and gender. This approach allowed us to test the unique association of each socioeconomic indicator with the learning parameters, while controlling for the other. In follow-up analyses, we examined the four subdomains of material hardship (food insecurity, clothing affordability, access to entertainment, and housing stability) in a separate regression model, with computational parameters as the dependent variable and household income, age, and gender included as covariates.</p>
</sec>
</sec>
<sec sec-type="results" id="sec7">
<label>3</label>
<title>Results</title>
<p>Participants completed 60 trials of the probabilistic reversal learning task, achieving a mean accuracy of 69.4% (<italic>SD</italic>&#x202F;=&#x202F;6.7%, <xref ref-type="table" rid="tab1">Table 1</xref>). A one-sample <italic>t</italic>-test confirmed that the overall accuracy (69.4%) was significantly above chance level (50%), <italic>t</italic>(94)&#x202F;=&#x202F;28.22, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001, Cohen&#x2019;s <italic>d</italic>&#x202F;=&#x202F;2.895, indicating successful learning throughout the task (<xref ref-type="fig" rid="fig2">Figure 2</xref> shows the trial-by-trial accuracy profile). The average number of successful reversals was 3.2 (<italic>SD</italic>&#x202F;=&#x202F;1.5). Computational modeling using a hierarchical Bayesian approach estimated individual parameters for reward learning rate (<italic>M</italic>&#x202F;=&#x202F;0.72, <italic>SD</italic>&#x202F;=&#x202F;0.01), punishment learning rate (<italic>M</italic>&#x202F;=&#x202F;0.54, <italic>SD</italic>&#x202F;=&#x202F;0.09), and inverse temperature (<italic>M</italic>&#x202F;=&#x202F;1.36, <italic>SD</italic>&#x202F;=&#x202F;0.62). Model convergence was confirmed by R-hat values&#x202F;&#x003C;&#x202F;1.01.</p>
<table-wrap position="float" id="tab1">
<label>Table 1</label>
<caption>
<p>Descriptive statistics and bivariate correlations among study variables.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Variable</th>
<th align="center" valign="top">1</th>
<th align="center" valign="top">2</th>
<th align="center" valign="top">3</th>
<th align="center" valign="top">4</th>
<th align="center" valign="top">5</th>
<th align="center" valign="top">6</th>
<th align="center" valign="top">7</th>
<th align="center" valign="top">8</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Mean</td>
<td align="center" valign="middle">8.14</td>
<td align="center" valign="middle">3.63</td>
<td align="center" valign="middle">69.4%</td>
<td align="center" valign="middle">603&#x202F;ms</td>
<td align="center" valign="middle">3.2</td>
<td align="center" valign="middle">0.719</td>
<td align="center" valign="middle">0.542</td>
<td align="center" valign="middle">1.362</td>
</tr>
<tr>
<td align="left" valign="middle">Standard deviation</td>
<td align="center" valign="middle">3.65</td>
<td align="center" valign="middle">1.62</td>
<td align="center" valign="middle">6.7%</td>
<td align="center" valign="middle">95&#x202F;ms</td>
<td align="center" valign="middle">1.5</td>
<td align="center" valign="middle">0.008</td>
<td align="center" valign="middle">0.089</td>
<td align="center" valign="middle">0.619</td>
</tr>
<tr>
<td align="left" valign="middle">1. Material hardship</td>
<td align="center" valign="middle">&#x2014;</td>
<td/>
<td/>
<td/>
<td/>
<td/>
<td/>
<td/>
</tr>
<tr>
<td align="left" valign="middle">2. Household income</td>
<td align="center" valign="middle">&#x2212;0.417&#x002A;&#x002A;&#x002A;</td>
<td align="center" valign="middle">&#x2014;</td>
<td/>
<td/>
<td/>
<td/>
<td/>
<td/>
</tr>
<tr>
<td align="left" valign="middle">3. Accuracy</td>
<td align="center" valign="middle">&#x2212;0.109</td>
<td align="center" valign="middle">&#x2212;0.041</td>
<td align="center" valign="middle">&#x2014;</td>
<td/>
<td/>
<td/>
<td/>
<td/>
</tr>
<tr>
<td align="left" valign="middle">4. Reaction time</td>
<td align="center" valign="middle">0.014</td>
<td align="center" valign="middle">&#x2212;0.004</td>
<td align="center" valign="middle">&#x2212;0.211&#x002A;</td>
<td align="center" valign="middle">&#x2014;</td>
<td/>
<td/>
<td/>
<td/>
</tr>
<tr>
<td align="left" valign="middle">5. Reversal frequency</td>
<td align="center" valign="middle">&#x2212;0.209&#x002A;</td>
<td align="center" valign="middle">0.125</td>
<td align="center" valign="middle">0.518&#x002A;&#x002A;&#x002A;</td>
<td align="center" valign="middle">&#x2212;0.086</td>
<td align="center" valign="middle">&#x2014;</td>
<td/>
<td/>
<td/>
</tr>
<tr>
<td align="left" valign="middle">6. Reward learning rate</td>
<td align="center" valign="middle">&#x2212;0.073</td>
<td align="center" valign="middle">0.103</td>
<td align="center" valign="middle">0.050</td>
<td align="center" valign="middle">&#x2212;0.048</td>
<td align="center" valign="middle">0.283&#x002A;&#x002A;</td>
<td align="center" valign="middle">&#x2014;</td>
<td/>
<td/>
</tr>
<tr>
<td align="left" valign="middle">7. Punishment learning rate</td>
<td align="center" valign="middle">0.261&#x002A;</td>
<td align="center" valign="middle">&#x2212;0.144</td>
<td align="center" valign="middle">0.298&#x002A;&#x002A;</td>
<td align="center" valign="middle">&#x2212;0.219&#x002A;</td>
<td align="center" valign="middle">&#x2212;0.352&#x002A;&#x002A;&#x002A;</td>
<td align="center" valign="middle">&#x2212;0.316&#x002A;&#x002A;</td>
<td align="center" valign="middle">&#x2014;</td>
<td/>
</tr>
<tr>
<td align="left" valign="middle">8. Inverse temperature</td>
<td align="center" valign="middle">&#x2212;0.077</td>
<td align="center" valign="middle">0.072</td>
<td align="center" valign="middle">0.579&#x002A;&#x002A;&#x002A;</td>
<td align="center" valign="middle">&#x2212;0.189</td>
<td align="center" valign="middle">0.513&#x002A;&#x002A;&#x002A;</td>
<td align="center" valign="middle">0.280&#x002A;&#x002A;</td>
<td align="center" valign="middle">0.085</td>
<td align="center" valign="middle">&#x2014;</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>&#x002A;<italic>p</italic>&#x202F;&#x003C;&#x202F;0.05, &#x002A;&#x002A;<italic>p</italic>&#x202F;&#x003C;&#x202F;0.01, &#x002A;&#x002A;&#x002A;<italic>p</italic>&#x202F;&#x003C;&#x202F;0.001.</p>
</table-wrap-foot>
</table-wrap>
<fig position="float" id="fig2">
<label>Figure 2</label>
<caption>
<p>Trial-by-Trial accuracy in reversal learning task.</p>
</caption>
<graphic xlink:href="fpsyg-16-1665380-g002.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Line chart plotting trial-by-trial accuracy over 60 trials. A solid blue line indicates mean accuracy, which fluctuates around 0.7, with a shaded area representing error margins. A dashed red line marks the 0.5 chance level.</alt-text>
</graphic>
</fig>
<p>Significant correlations emerged between task performance and model parameters: reward learning rate was positively associated with reversal frequency (<italic>r</italic>&#x202F;=&#x202F;0.283, <italic>p</italic>&#x202F;=&#x202F;0.006, 95% CI [0.089, 0.457]) but not significantly associated with accuracy (<italic>r</italic>&#x202F;=&#x202F;0.050, <italic>p</italic>&#x202F;=&#x202F;0.632, 95% CI [&#x2212;0.124, 0.224]). While punishment learning rate showed positive correlation with accuracy (<italic>r</italic>&#x202F;=&#x202F;0.298, <italic>p</italic>&#x202F;=&#x202F;0.003, 95% CI [0.133, 0.438]), it was negatively associated with reversal frequency (<italic>r</italic>&#x202F;=&#x202F;&#x2212;0.352, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001, 95% CI [&#x2212;0.511, &#x2212;0.180]). The inverse temperature parameter positively correlated with both accuracy (<italic>r</italic>&#x202F;=&#x202F;0.579, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001, 95% CI [0.441, 0.699]) and reversal frequency (<italic>r</italic>&#x202F;=&#x202F;0.513, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001, 95% CI [0.362, 0.649]).</p>
<p>Bivariate analyses revealed that material hardship (<italic>M</italic>&#x202F;=&#x202F;8.14, <italic>SD</italic>&#x202F;=&#x202F;3.65) and household income (<italic>M</italic>&#x202F;=&#x202F;3.63, <italic>SD</italic>&#x202F;=&#x202F;1.62) were inversely correlated (<italic>r</italic>&#x202F;=&#x202F;&#x2212;0.417, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001, 95% CI [&#x2212;0.554, &#x2212;0.251]). Material hardship correlated positively with punishment learning rate (<italic>r</italic>&#x202F;=&#x202F;0.261, <italic>p</italic>&#x202F;=&#x202F;0.011, 95% CI [0.080, 0.435], <xref ref-type="fig" rid="fig3">Figure 3</xref>) and negatively with reversal frequency (<italic>r</italic>&#x202F;=&#x202F;&#x2212;0.209, <italic>p</italic>&#x202F;=&#x202F;0.042, 95% CI [&#x2212;0.382, &#x2212;0.032]). Household income showed no significant correlations with reward learning rate (<italic>r</italic>&#x202F;=&#x202F;0.103, <italic>p</italic>&#x202F;=&#x202F;0.323, 95% CI [&#x2212;0.068, 0.273]), punishment learning rate (<italic>r</italic>&#x202F;=&#x202F;&#x2212;0.144, <italic>p</italic>&#x202F;=&#x202F;0.164, 95% CI [&#x2212;0.330, 0.042]), or inverse temperature (<italic>r</italic>&#x202F;=&#x202F;0.072, <italic>p</italic>&#x202F;=&#x202F;0.490, 95% CI [&#x2212;0.125, 0.277]). Among hardship subdomains, housing instability (<italic>M</italic>&#x202F;=&#x202F;2.38, <italic>SD</italic>&#x202F;=&#x202F;1.40) showed the strongest correlation with reversal frequency (<italic>r</italic>&#x202F;=&#x202F;&#x2212;0.218, <italic>p</italic>&#x202F;=&#x202F;0.034, 95% CI [&#x2212;0.408, &#x2212;0.014]) and punishment learning rate (<italic>r</italic>&#x202F;=&#x202F;0.274, <italic>p</italic>&#x202F;=&#x202F;0.007, 95% CI [0.062, 0.460]).</p>
<fig position="float" id="fig3">
<label>Figure 3</label>
<caption>
<p>Partial association between material hardship and punishment learning rate after controlling for age, gender, and household income.</p>
</caption>
<graphic xlink:href="fpsyg-16-1665380-g003.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Scatter plot displaying the partial association between Material Hardship and Punishment Learning Rate, after controlling for covariates. Data points are distributed along a weakly positive trend line, with dashed curves indicating the confidence interval.</alt-text>
</graphic>
</fig>
<p>Multiple regression analyses, which included both socioeconomic indicators while controlling for age and gender, revealed a distinct pattern of associations. Of primary theoretical interest, material hardship showed a positive association with punishment learning rate (<italic>b</italic>&#x202F;=&#x202F;0.240, 95% CI [0.016, 0.464], uncorrected <italic>p</italic>&#x202F;=&#x202F;0.036, FDR-corrected <italic>p</italic>&#x202F;=&#x202F;0.108, <xref ref-type="table" rid="tab2">Table 2</xref>). Although this association did not survive FDR correction, the medium effect size and the confidence interval excluding zero suggest a meaningful pattern consistent with the hypothesis that economic hardship sensitizes individuals to negative outcomes. In contrast, household income was not meaningfully associated with punishment learning rate (<italic>b</italic>&#x202F;=&#x202F;&#x2212;0.046, <italic>p</italic>&#x202F;=&#x202F;0.685) or any other model parameters (reward learning rate: <italic>b</italic>&#x202F;=&#x202F;0.090, <italic>p</italic>&#x202F;=&#x202F;0.439; inverse temperature: <italic>b</italic>&#x202F;=&#x202F;0.038, <italic>p</italic>&#x202F;=&#x202F;0.739). Material hardship itself demonstrated specificity, as it was not associated with reward learning rate (<italic>b</italic>&#x202F;=&#x202F;&#x2212;0.034, <italic>p</italic>&#x202F;=&#x202F;0.773) or inverse temperature (<italic>b</italic>&#x202F;=&#x202F;&#x2212;0.067, <italic>p</italic>&#x202F;=&#x202F;0.560). Neither age nor gender predicted any learning parameters (all <italic>p</italic>&#x202F;&#x003E;&#x202F;0.05). In exploratory follow-up regression models that examined hardship subdomains individually, only housing instability emerged as the unique predictor of punishment learning rate (<italic>b</italic>&#x202F;=&#x202F;0.255, <italic>p</italic>&#x202F;=&#x202F;0.022). Multicollinearity diagnostics indicated no concerns (all variance inflation factors&#x202F;&#x003C;&#x202F;1.3).</p>
<table-wrap position="float" id="tab2">
<label>Table 2</label>
<caption>
<p>Multiple regression analyses predicting reinforcement learning parameters.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top" rowspan="2">Predictors</th>
<th align="center" valign="top" colspan="3">Reward learning rate</th>
<th align="center" valign="top" colspan="3">Punishment learning rate</th>
<th align="center" valign="top" colspan="3">Inverse temperature</th>
</tr>
<tr>
<th align="center" valign="top"><italic>b</italic></th>
<th align="center" valign="top">95% CI</th>
<th align="center" valign="top"><italic>p</italic></th>
<th align="center" valign="top"><italic>b</italic></th>
<th align="center" valign="top">95% CI</th>
<th align="center" valign="top"><italic>p</italic></th>
<th align="center" valign="top"><italic>b</italic></th>
<th align="center" valign="top">95% CI</th>
<th align="center" valign="top"><italic>p</italic></th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Material hardship</td>
<td align="center" valign="middle">&#x2212;0.034</td>
<td align="center" valign="middle">[&#x2212;0.263, 0.196]</td>
<td align="center" valign="middle">0.773</td>
<td align="center" valign="middle">0.240</td>
<td align="center" valign="middle">[0.016, 0.464]</td>
<td align="center" valign="middle">0.036</td>
<td align="center" valign="middle">&#x2212;0.067</td>
<td align="center" valign="middle">[&#x2212;0.296, 0.161]</td>
<td align="center" valign="middle">0.560</td>
</tr>
<tr>
<td align="left" valign="middle">Household income</td>
<td align="center" valign="middle">0.090</td>
<td align="center" valign="middle">[&#x2212;0.140, 0.319]</td>
<td align="center" valign="middle">0.439</td>
<td align="center" valign="middle">&#x2212;0.046</td>
<td align="center" valign="middle">[&#x2212;0.270, 0.178]</td>
<td align="center" valign="middle">0.685</td>
<td align="center" valign="middle">0.038</td>
<td align="center" valign="middle">[&#x2212;0.190, 0.267]</td>
<td align="center" valign="middle">0.739</td>
</tr>
<tr>
<td align="left" valign="middle">Age</td>
<td align="center" valign="middle">&#x2212;0.045</td>
<td align="center" valign="middle">[&#x2212;0.258, 0.167]</td>
<td align="center" valign="middle">0.673</td>
<td align="center" valign="middle">0.005</td>
<td align="center" valign="middle">[&#x2212;0.203, 0.212]</td>
<td align="center" valign="middle">0.965</td>
<td align="center" valign="middle">0.014</td>
<td align="center" valign="middle">[&#x2212;0.198, 0.226]</td>
<td align="center" valign="middle">0.895</td>
</tr>
<tr>
<td align="left" valign="middle">Gender</td>
<td align="center" valign="middle">0.114</td>
<td align="center" valign="middle">[&#x2212;0.098, 0.325]</td>
<td align="center" valign="middle">0.288</td>
<td align="center" valign="middle">&#x2212;0.051</td>
<td align="center" valign="middle">[&#x2212;0.257, 0.155]</td>
<td align="center" valign="middle">0.624</td>
<td align="center" valign="middle">&#x2212;0.155</td>
<td align="center" valign="middle">[&#x2212;0.365, 0.056]</td>
<td align="center" valign="middle">0.148</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p><italic>b</italic>&#x202F;=&#x202F;standardized beta coefficient. CI, confidence interval.</p>
</table-wrap-foot>
</table-wrap>
</sec>
<sec sec-type="discussion" id="sec8">
<label>4</label>
<title>Discussion</title>
<p>This study directly addresses the critical gap concerning how socioeconomic disadvantage shapes RL mechanisms. By implementing a probabilistic reversal learning paradigm with a computational model that dissociates three core parameters&#x2014;reward learning rate, punishment learning rate, and inverse temperature&#x2014;we systematically evaluated the unique contributions of material hardship versus household income. Our findings reveal a targeted learning impairment: individuals experiencing material hardship, characterized by direct deprivation of basic needs, specifically exhibit heightened behavioral responsiveness to negative outcomes while maintaining intact reward processing. In contrast, household income demonstrated no significant relationship with any learning parameter. Among specific hardship subtypes, housing instability emerged as the strongest driver of this punishment sensitivity effect. Collectively, these results demonstrate how immediate deprivation experiences reconfigure fundamental learning mechanisms independently of financial constraints.</p>
<p>Our findings suggest a potential dissociation between socioeconomic indicators. Specifically, material hardship was associated with an elevated punishment learning rate, indicating heightened sensitivity to negative feedback, while reward learning remained unaffected. Although this association should be interpreted with caution as it did not survive strict correction for multiple comparisons, the observed effect size suggests a pattern worthy of further investigation. We propose that this hypersensitivity to punishment could become maladaptive in the current task by directly obstructing the acquisition of the latent task structure. Specifically, during learning phases which require ignoring occasional negative feedback to persist with the correct option, excessive reactivity to punishments causes premature abandonment of advantageous choices. Rather than tolerating probabilistic losses to maintain correct responding, they over-interpret negative outcomes as signals to switch strategies. This pattern reflects a failure to integrate feedback in a context-appropriate manner, ultimately obstructing the learning of latent task structure. The tendency to prioritize reactive switching over stable goal-directed behavior aligns with previous accounts of how adversity can bias decision-making under uncertainty (<xref ref-type="bibr" rid="ref23">Lisi et al., 2025</xref>). Thus, socioeconomic disadvantage may recalibrate cognitive processes toward heightened reactivity to negative outcomes, perpetuating disadvantage cycles through maladaptive behavioral patterns.</p>
<p>Critically, our analyses demonstrate that material hardship, not household income, is the decisive socioeconomic factor driving alterations in punishment learning. While household income and material hardship are closely correlated, material hardship uniquely predicted both heightened punishment learning rates and poorer behavioral adaptation (i.e., reduced reversals). This dissociation aligns with longitudinal evidence showing material hardship independently predicts cognitive deficits beyond income effects (<xref ref-type="bibr" rid="ref8">Daniel et al., 2024</xref>). We propose this occurs because immediate hardship generates perceived stress (<xref ref-type="bibr" rid="ref19">Huang et al., 2021</xref>), which disproportionately overburdens neurocognitive systems governing threat response. Consequently, individuals become hyper-responsive to losses at the expense of adaptive flexibility. This pattern supports theoretical frameworks positing that distal variables shape the current life situation (<xref ref-type="bibr" rid="ref27">Martin and Martin, 2002</xref>). Future research should prioritize measuring direct adversity experiences&#x2014;such as unstable housing&#x2014;as the critical pathways connecting socioeconomic disadvantage to cognitive changes.</p>
<p>Examining the subdomains of material hardship more closely, we found that housing instability emerged as the strongest predictor of impaired punishment learning. Longitudinal research showed that housing instability had a stronger effect on cognitive development than child maltreatment, poverty, and other risks (<xref ref-type="bibr" rid="ref13">Fowler et al., 2015</xref>). A scoping review highlighted cognitive impairment as both a risk factor for and a consequence of homelessness (<xref ref-type="bibr" rid="ref32">Stone et al., 2019</xref>). Unlike other financial pressures, housing insecurity uniquely compromises fundamental safety needs, keeping individuals in survival-mode cognition. Neuroimaging evidence confirms such adversity amplifies amygdala reactivity to stress (<xref ref-type="bibr" rid="ref36">Tottenham, 2009</xref>), which partly explaining our findings. Critically, interventions that stabilizing housing&#x2014;like housing vouchers&#x2014;showed measurable psychological benefits (<xref ref-type="bibr" rid="ref12">Finnie et al., 2022</xref>), making housing stability interventions a highly effective policy approach to reduce harmful cognitive effects linked to socioeconomic disadvantage.</p>
<p>Our identification of material hardship&#x2014;particularly housing instability&#x2014;as a primary mechanism driving maladaptive punishment learning necessitates structural policy interventions. Critically, approaches focused exclusively on income supplementation might be less effective in addressing the cognitive consequences of direct deprivation experiences. Effective solutions must instead target the tangible manifestations of material hardship through comprehensive social safety nets. These should include: (1) housing stabilization programs with eviction protection, (2) expanded food assistance, (3) universal healthcare access, and (4) guaranteed utility support. Such interventions directly reduce the chronic stress and perceived scarcity stemming from unmet basic needs&#x2014;precisely the mechanism through which hardship amplifies neural sensitivity to negative outcomes in our study. By ensuring environmental stability, these policies create conditions conducive to neurocognitive recovery. As our findings demonstrate that secure housing specifically mitigates punishment hypersensitivity, prioritizing these multi-faceted supports will foster improved learning flexibility and adaptive decision-making in disadvantaged communities, ultimately disrupting cycles of socioeconomic disadvantage.</p>
<p>While this study advances our understanding of how socioeconomic disadvantage shapes learning, several limitations should be acknowledged. First, the cross-sectional design limits causal inference. Future longitudinal research should track how socioeconomic disadvantage influence learning mechanisms across development and examine whether interventions can modify these pathways. Second, our participant sample limits generalizability. Replication studies with more diverse populations and age groups are needed, particularly in understanding how socioeconomic disadvantage affects neurodevelopment in children and adolescents. Third, our computational modeling approach employed a parsimonious three-parameter model that dissociates reward and punishment learning rates. While this model choice was appropriate for our sample size, it does not capture all aspects of reinforcement learning, such as separate scaling parameters for reward and punishment sensitivity. Future studies with larger samples could employ more complex models to provide a more comprehensive account. Fourth, although computational modeling provides relatively precise parameter estimates, it cannot fully capture complex cognitive processes. Future work should therefore combine computational modeling with neuroimaging techniques, in order to map model-derived cognitive processes to their neural substrates and identify how hardship exposure affects neural systems underlying RL. Finally, future intervention research should empirically evaluate cognitive and behavioral strategies specifically designed to mitigate maladaptive patterns in punishment learning among disadvantaged populations&#x2014;strategies that directly target the neurocognitive mechanisms identified here&#x2014;with rigorous measurement of their efficacy in disrupting cycles of socioeconomic disadvantage.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="sec9">
<title>Data availability statement</title>
<p>The datasets presented in this study can be found in online repositories. The names of the repository/repositories and accession number(s) can be found at: <ext-link xlink:href="https://osf.io/zchxy/?view_only=88519feb52e244298b3c2dd0b225d8f0%3c/b%3e" ext-link-type="uri">https://osf.io/zchxy/?view_only=88519feb52e244298b3c2dd0b225d8f0</ext-link>.</p>
</sec>
<sec sec-type="ethics-statement" id="sec10">
<title>Ethics statement</title>
<p>The studies involving humans were approved by the Institutional Review Board of South China Normal University. The studies were conducted in accordance with the local legislation and institutional requirements. The participants provided their written informed consent to participate in this study.</p>
</sec>
<sec sec-type="author-contributions" id="sec11">
<title>Author contributions</title>
<p>ZW: Writing &#x2013; original draft, Conceptualization, Formal analysis, Investigation, Validation. XH: Formal analysis, Investigation, Methodology, Writing &#x2013; review &#x0026; editing. YS: Investigation, Validation, Writing &#x2013; review &#x0026; editing. LB: Investigation, Validation, Writing &#x2013; review &#x0026; editing. YW: Conceptualization, Funding acquisition, Supervision, Writing &#x2013; review &#x0026; editing.</p>
</sec>
<sec sec-type="funding-information" id="sec12">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research and/or publication of this article. This research was supported by the National Social Science Fund of China (grant no. 24BXW075), the Guangdong Philosophy and Social Sciences Planning Project 2023 (grant no. GD23CMK04), and the Collaborative Center for the Promotion of Chinese Culture in Hong Kong, Macau, Taiwan, and Overseas at Jinan University (grant no. JNXT2023002).</p>
</sec>
<sec sec-type="COI-statement" id="sec13">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="sec14">
<title>Generative AI statement</title>
<p>The authors declare that no Gen AI was used in the creation of this manuscript.</p>
<p>Any alternative text (alt text) provided alongside figures in this article has been generated by Frontiers with the support of artificial intelligence and reasonable efforts have been made to ensure accuracy, including review by the authors wherever possible. If you identify any issues, please contact us.</p>
</sec>
<sec sec-type="disclaimer" id="sec15">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="ref1"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ahn</surname><given-names>W.-Y.</given-names></name> <name><surname>Haines</surname><given-names>N.</given-names></name> <name><surname>Zhang</surname><given-names>L.</given-names></name></person-group> (<year>2017</year>). <article-title>Revealing neurocomputational mechanisms of reinforcement learning and decision-making with the hBayesDM package</article-title>. <source>Comput. Psychiatry</source> <volume>1</volume>:<fpage>24</fpage>. doi: <pub-id pub-id-type="doi">10.1162/CPSY_a_00002</pub-id>, PMID: <pub-id pub-id-type="pmid">29601060</pub-id></citation></ref>
<ref id="ref2"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Atombo</surname><given-names>C.</given-names></name> <name><surname>Wu</surname><given-names>C.</given-names></name> <name><surname>Tettehfio</surname><given-names>E. O.</given-names></name> <name><surname>Agbo</surname><given-names>A. A.</given-names></name></person-group> (<year>2017</year>). <article-title>Personality, socioeconomic status, attitude, intention and risky driving behavior</article-title>. <source>Cogent Psychol.</source> <volume>4</volume>:<fpage>1376424</fpage>. doi: <pub-id pub-id-type="doi">10.1080/23311908.2017.1376424</pub-id></citation></ref>
<ref id="ref3"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Benjamini</surname><given-names>Y.</given-names></name> <name><surname>Hochberg</surname><given-names>Y.</given-names></name></person-group> (<year>1995</year>). <article-title>Controlling the false discovery rate: a practical and powerful approach to multiple testing</article-title>. <source>J. R. Stat. Soc. Ser. B Stat Methodol.</source> <volume>57</volume>, <fpage>289</fpage>&#x2013;<lpage>300</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.2517-6161.1995.tb02031.x</pub-id></citation></ref>
<ref id="ref4"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Beverly</surname><given-names>S. G.</given-names></name></person-group> (<year>2001</year>). <article-title>Measures of material hardship: rationale and recommendations</article-title>. <source>J. Poverty</source> <volume>5</volume>, <fpage>23</fpage>&#x2013;<lpage>41</lpage>. doi: <pub-id pub-id-type="doi">10.1300/J134v05n01_02</pub-id></citation></ref>
<ref id="ref5"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chen</surname><given-names>C.</given-names></name> <name><surname>Wang</surname><given-names>Z.</given-names></name> <name><surname>Cao</surname><given-names>X.</given-names></name> <name><surname>Zhu</surname><given-names>J.</given-names></name></person-group> (<year>2023</year>). <article-title>Exploring the association between early exposure to material hardship and psychopathology through indirect effects of fronto-limbic functional connectivity during fear learning</article-title>. <source>Cereb. Cortex</source> <volume>33</volume>, <fpage>10702</fpage>&#x2013;<lpage>10710</lpage>. doi: <pub-id pub-id-type="doi">10.1093/cercor/bhad320</pub-id>, PMID: <pub-id pub-id-type="pmid">37689831</pub-id></citation></ref>
<ref id="ref6"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Daniel</surname><given-names>R.</given-names></name> <name><surname>Pollmann</surname><given-names>S.</given-names></name></person-group> (<year>2014</year>). <article-title>A universal role of the ventral striatum in reward-based learning: evidence from human studies</article-title>. <source>Neurobiol. Learn. Mem.</source> <volume>114</volume>, <fpage>90</fpage>&#x2013;<lpage>100</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.nlm.2014.05.002</pub-id>, PMID: <pub-id pub-id-type="pmid">24825620</pub-id></citation></ref>
<ref id="ref7"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Daniel</surname><given-names>R.</given-names></name> <name><surname>Radulescu</surname><given-names>A.</given-names></name> <name><surname>Niv</surname><given-names>Y.</given-names></name></person-group> (<year>2020</year>). <article-title>Intact reinforcement learning but impaired attentional control during multidimensional probabilistic learning in older adults</article-title>. <source>J. Neurosci.</source> <volume>40</volume>, <fpage>1084</fpage>&#x2013;<lpage>1096</lpage>. doi: <pub-id pub-id-type="doi">10.1523/JNEUROSCI.0254-19.2019</pub-id>, PMID: <pub-id pub-id-type="pmid">31826943</pub-id></citation></ref>
<ref id="ref8"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Daniel</surname><given-names>G.</given-names></name> <name><surname>Williams</surname><given-names>C.</given-names></name> <name><surname>Lawrence</surname><given-names>A.</given-names></name> <name><surname>Buckley</surname><given-names>K.</given-names></name> <name><surname>Leonard</surname><given-names>D.</given-names></name> <name><surname>Bernal</surname><given-names>D.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>Income-based poverty and material hardship predict reduced cognitive performance in older American adults</article-title>. <source>Innov. Aging</source> <volume>8</volume>:<fpage>1322</fpage>. doi: <pub-id pub-id-type="doi">10.1093/geroni/igae098.4221</pub-id></citation></ref>
<ref id="ref9"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>De Bruijn</surname><given-names>E.-J.</given-names></name> <name><surname>Antonides</surname><given-names>G.</given-names></name></person-group> (<year>2022</year>). <article-title>Poverty and economic decision making: a review of scarcity theory</article-title>. <source>Theor. Decis.</source> <volume>92</volume>, <fpage>5</fpage>&#x2013;<lpage>37</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s11238-021-09802-7</pub-id></citation></ref>
<ref id="ref10"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>den Ouden</surname><given-names>H. E. M.</given-names></name> <name><surname>Daw</surname><given-names>N. D.</given-names></name> <name><surname>Fernandez</surname><given-names>G.</given-names></name> <name><surname>Elshout</surname><given-names>J. A.</given-names></name> <name><surname>Rijpkema</surname><given-names>M.</given-names></name> <name><surname>Hoogman</surname><given-names>M.</given-names></name> <etal/></person-group>. (<year>2013</year>). <article-title>Dissociable effects of dopamine and serotonin on reversal learning</article-title>. <source>Neuron</source> <volume>80</volume>, <fpage>1090</fpage>&#x2013;<lpage>1100</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.neuron.2013.08.030</pub-id>, PMID: <pub-id pub-id-type="pmid">24267657</pub-id></citation></ref>
<ref id="ref11"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Faul</surname><given-names>F.</given-names></name> <name><surname>Erdfelder</surname><given-names>E.</given-names></name> <name><surname>Lang</surname><given-names>A.-G.</given-names></name> <name><surname>Buchner</surname><given-names>A.</given-names></name></person-group> (<year>2007</year>). <article-title>G&#x002A;power 3: a flexible statistical power analysis program for the social, behavioral, and biomedical sciences</article-title>. <source>Behav. Res. Methods</source> <volume>39</volume>, <fpage>175</fpage>&#x2013;<lpage>191</lpage>. doi: <pub-id pub-id-type="doi">10.3758/BF03193146</pub-id>, PMID: <pub-id pub-id-type="pmid">17695343</pub-id></citation></ref>
<ref id="ref12"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Finnie</surname><given-names>R. K. C.</given-names></name> <name><surname>Peng</surname><given-names>Y.</given-names></name> <name><surname>Hahn</surname><given-names>R. A.</given-names></name> <name><surname>Schwartz</surname><given-names>A.</given-names></name> <name><surname>Emmons</surname><given-names>K.</given-names></name> <name><surname>Montgomery</surname><given-names>A. E.</given-names></name> <etal/></person-group>. (<year>2022</year>). <article-title>Tenant-based housing voucher programs: a community guide systematic review</article-title>. <source>J. Public Health Manag. Pract.</source> <volume>28</volume>, <fpage>E795</fpage>&#x2013;<lpage>E803</lpage>. doi: <pub-id pub-id-type="doi">10.1097/phh.0000000000001588</pub-id>, PMID: <pub-id pub-id-type="pmid">36194822</pub-id></citation></ref>
<ref id="ref13"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fowler</surname><given-names>P. J.</given-names></name> <name><surname>McGrath</surname><given-names>L. M.</given-names></name> <name><surname>Henry</surname><given-names>D. B.</given-names></name> <name><surname>Schoeny</surname><given-names>M.</given-names></name> <name><surname>Chavira</surname><given-names>D.</given-names></name> <name><surname>Taylor</surname><given-names>J. J.</given-names></name> <etal/></person-group>. (<year>2015</year>). <article-title>Housing mobility and cognitive development: change in verbal and nonverbal abilities</article-title>. <source>Child Abuse Negl.</source> <volume>48</volume>, <fpage>104</fpage>&#x2013;<lpage>118</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.chiabu.2015.06.002</pub-id>, PMID: <pub-id pub-id-type="pmid">26184055</pub-id></citation></ref>
<ref id="ref14"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gershoff</surname><given-names>E. T.</given-names></name> <name><surname>Aber</surname><given-names>J. L.</given-names></name> <name><surname>Raver</surname><given-names>C. C.</given-names></name> <name><surname>Lennon</surname><given-names>M. C.</given-names></name></person-group> (<year>2007</year>). <article-title>Income is not enough: incorporating material hardship into models of income associations with parenting and child development</article-title>. <source>Child Dev.</source> <volume>78</volume>, <fpage>70</fpage>&#x2013;<lpage>95</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1467-8624.2007.00986.x</pub-id>, PMID: <pub-id pub-id-type="pmid">17328694</pub-id></citation></ref>
<ref id="ref15"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gl&#x00E4;scher</surname><given-names>J.</given-names></name> <name><surname>Hampton</surname><given-names>A. N.</given-names></name> <name><surname>O&#x2019;Doherty</surname><given-names>J. P.</given-names></name></person-group> (<year>2009</year>). <article-title>Determining a role for ventromedial prefrontal cortex in encoding action-based value signals during reward-related decision making</article-title>. <source>Cereb. Cortex</source> <volume>19</volume>, <fpage>483</fpage>&#x2013;<lpage>495</lpage>. doi: <pub-id pub-id-type="doi">10.1093/cercor/bhn098</pub-id></citation></ref>
<ref id="ref16"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gueguen</surname><given-names>M. C. M.</given-names></name> <name><surname>Lopez-Persem</surname><given-names>A.</given-names></name> <name><surname>Billeke</surname><given-names>P.</given-names></name> <name><surname>Lachaux</surname><given-names>J.-P.</given-names></name> <name><surname>Rheims</surname><given-names>S.</given-names></name> <name><surname>Kahane</surname><given-names>P.</given-names></name> <etal/></person-group>. (<year>2021</year>). <article-title>Anatomical dissociation of intracerebral signals for reward and punishment prediction errors in humans</article-title>. <source>Nat. Commun.</source> <volume>12</volume>:<fpage>3344</fpage>. doi: <pub-id pub-id-type="doi">10.1038/s41467-021-23704-w</pub-id>, PMID: <pub-id pub-id-type="pmid">34099678</pub-id></citation></ref>
<ref id="ref17"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>He</surname><given-names>X.</given-names></name> <name><surname>Qiu</surname><given-names>B.</given-names></name> <name><surname>Deng</surname><given-names>Y.</given-names></name> <name><surname>Wang</surname><given-names>Z.</given-names></name> <name><surname>Cao</surname><given-names>X.</given-names></name> <name><surname>Zheng</surname><given-names>X.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>Material hardship predicts response bias in loss-averse decisions: the roles of anxiety and cognitive control</article-title>. <source>J. Psychol.</source> <volume>158</volume>, <fpage>309</fpage>&#x2013;<lpage>324</lpage>. doi: <pub-id pub-id-type="doi">10.1080/00223980.2023.2296946</pub-id>, PMID: <pub-id pub-id-type="pmid">38227200</pub-id></citation></ref>
<ref id="ref18"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Heininga</surname><given-names>V. E.</given-names></name> <name><surname>Van Roekel</surname><given-names>E.</given-names></name> <name><surname>Wichers</surname><given-names>M.</given-names></name> <name><surname>Oldehinkel</surname><given-names>A. J.</given-names></name></person-group> (<year>2017</year>). <article-title>Reward and punishment learning in daily life: a replication study</article-title>. <source>PLoS One</source> <volume>12</volume>:<fpage>e0180753</fpage>. doi: <pub-id pub-id-type="doi">10.1371/journal.pone.0180753</pub-id>, PMID: <pub-id pub-id-type="pmid">28976985</pub-id></citation></ref>
<ref id="ref19"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Huang</surname><given-names>Y.</given-names></name> <name><surname>Heflin</surname><given-names>C. M.</given-names></name> <name><surname>Validova</surname><given-names>A.</given-names></name></person-group> (<year>2021</year>). <article-title>Material hardship, perceived stress, and health in early adulthood</article-title>. <source>Ann. Epidemiol.</source> <volume>53</volume>, <fpage>69</fpage>&#x2013;<lpage>75.e3</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.annepidem.2020.08.017</pub-id>, PMID: <pub-id pub-id-type="pmid">32949721</pub-id></citation></ref>
<ref id="ref20"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Katahira</surname><given-names>K.</given-names></name></person-group> (<year>2015</year>). <article-title>The relation between reinforcement learning parameters and the influence of reinforcement history on choice behavior</article-title>. <source>J. Math. Psychol.</source> <volume>66</volume>, <fpage>59</fpage>&#x2013;<lpage>69</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.jmp.2015.03.006</pub-id></citation></ref>
<ref id="ref21"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Koch</surname><given-names>K.</given-names></name> <name><surname>Schachtzabel</surname><given-names>C.</given-names></name> <name><surname>Wagner</surname><given-names>G.</given-names></name> <name><surname>Reichenbach</surname><given-names>J. R.</given-names></name> <name><surname>Sauer</surname><given-names>H.</given-names></name> <name><surname>Schl&#x00F6;sser</surname><given-names>R.</given-names></name></person-group> (<year>2008</year>). <article-title>The neural correlates of reward-related trial-and-error learning: an fMRI study with a probabilistic learning task</article-title>. <source>Learn. Mem.</source> <volume>15</volume>, <fpage>728</fpage>&#x2013;<lpage>732</lpage>. doi: <pub-id pub-id-type="doi">10.1101/lm.1106408</pub-id>, PMID: <pub-id pub-id-type="pmid">18832559</pub-id></citation></ref>
<ref id="ref22"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lee</surname><given-names>D.</given-names></name> <name><surname>Seo</surname><given-names>H.</given-names></name> <name><surname>Jung</surname><given-names>M. W.</given-names></name></person-group> (<year>2012</year>). <article-title>Neural basis of reinforcement learning and decision making</article-title>. <source>Annu. Rev. Neurosci.</source> <volume>35</volume>, <fpage>287</fpage>&#x2013;<lpage>308</lpage>. doi: <pub-id pub-id-type="doi">10.1146/annurev-neuro-062111-150512</pub-id>, PMID: <pub-id pub-id-type="pmid">22462543</pub-id></citation></ref>
<ref id="ref23"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lisi</surname><given-names>M.</given-names></name> <name><surname>Michalek</surname><given-names>J.</given-names></name> <name><surname>Hadfield</surname><given-names>K.</given-names></name> <name><surname>Dajani</surname><given-names>R.</given-names></name> <name><surname>Mareschal</surname><given-names>I.</given-names></name></person-group> (<year>2025</year>). <article-title>Effects of early adversity and war trauma on learning under uncertainty</article-title>. <source>Dev. Sci.</source> <volume>28</volume>:<fpage>e70049</fpage>. doi: <pub-id pub-id-type="doi">10.1111/desc.70049</pub-id>, PMID: <pub-id pub-id-type="pmid">40778529</pub-id></citation></ref>
<ref id="ref24"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Mani</surname><given-names>A.</given-names></name> <name><surname>Mullainathan</surname><given-names>S.</given-names></name> <name><surname>Shafir</surname><given-names>E.</given-names></name> <name><surname>Zhao</surname><given-names>J.</given-names></name></person-group> (<year>2013</year>). <article-title>Poverty impedes cognitive function</article-title>. <source>Science</source> <volume>341</volume>, <fpage>976</fpage>&#x2013;<lpage>980</lpage>. doi: <pub-id pub-id-type="doi">10.1126/science.1238041</pub-id>, PMID: <pub-id pub-id-type="pmid">23990553</pub-id></citation></ref>
<ref id="ref25"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marbin</surname><given-names>D.</given-names></name> <name><surname>Gutwinski</surname><given-names>S.</given-names></name> <name><surname>Schreiter</surname><given-names>S.</given-names></name> <name><surname>Heinz</surname><given-names>A.</given-names></name></person-group> (<year>2022</year>). <article-title>Perspectives in poverty and mental health</article-title>. <source>Front. Public Health</source> <volume>10</volume>:<fpage>975482</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpubh.2022.975482</pub-id>, PMID: <pub-id pub-id-type="pmid">35991010</pub-id></citation></ref>
<ref id="ref26"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Marshall</surname><given-names>A. T.</given-names></name> <name><surname>Kirkpatrick</surname><given-names>K.</given-names></name></person-group> (<year>2017</year>). <article-title>Reinforcement learning models of risky choice and the promotion of risk-taking by losses disguised as wins in rats</article-title>. <source>J. Exp. Psychol. Anim. Learn. Cogn.</source> <volume>43</volume>, <fpage>262</fpage>&#x2013;<lpage>279</lpage>. doi: <pub-id pub-id-type="doi">10.1037/xan0000141</pub-id>, PMID: <pub-id pub-id-type="pmid">29120214</pub-id></citation></ref>
<ref id="ref27"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Martin</surname><given-names>P.</given-names></name> <name><surname>Martin</surname><given-names>M.</given-names></name></person-group> (<year>2002</year>). <article-title>Proximal and distal influences on development: the model of developmental adaptation</article-title>. <source>Dev. Rev.</source> <volume>22</volume>, <fpage>78</fpage>&#x2013;<lpage>96</lpage>. doi: <pub-id pub-id-type="doi">10.1006/drev.2001.0538</pub-id></citation></ref>
<ref id="ref28"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Romens</surname><given-names>S. E.</given-names></name> <name><surname>Casement</surname><given-names>M. D.</given-names></name> <name><surname>McAloon</surname><given-names>R.</given-names></name> <name><surname>Keenan</surname><given-names>K.</given-names></name> <name><surname>Hipwell</surname><given-names>A. E.</given-names></name> <name><surname>Guyer</surname><given-names>A. E.</given-names></name> <etal/></person-group>. (<year>2015</year>). <article-title>Adolescent girls&#x2019; neural response to reward mediates the relation between childhood financial disadvantage and depression</article-title>. <source>Child Psychol. Psychiatry</source> <volume>56</volume>, <fpage>1177</fpage>&#x2013;<lpage>1184</lpage>. doi: <pub-id pub-id-type="doi">10.1111/jcpp.12410</pub-id>, PMID: <pub-id pub-id-type="pmid">25846746</pub-id></citation></ref>
<ref id="ref29"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schaaf</surname><given-names>J. V.</given-names></name> <name><surname>Weidinger</surname><given-names>L.</given-names></name> <name><surname>Molleman</surname><given-names>L.</given-names></name> <name><surname>Van Den Bos</surname><given-names>W.</given-names></name></person-group> (<year>2023</year>). <article-title>Test&#x2013;retest reliability of reinforcement learning parameters</article-title>. <source>Behav. Res.</source> <volume>56</volume>, <fpage>4582</fpage>&#x2013;<lpage>4599</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13428-023-02203-4</pub-id>, PMID: <pub-id pub-id-type="pmid">37684495</pub-id></citation></ref>
<ref id="ref30"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shteingart</surname><given-names>H.</given-names></name> <name><surname>Loewenstein</surname><given-names>Y.</given-names></name></person-group> (<year>2014</year>). <article-title>Reinforcement learning and human behavior</article-title>. <source>Curr. Opin. Neurobiol.</source> <volume>25</volume>, <fpage>93</fpage>&#x2013;<lpage>98</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.conb.2013.12.004</pub-id>, PMID: <pub-id pub-id-type="pmid">24709606</pub-id></citation></ref>
<ref id="ref31"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Soltani</surname><given-names>A.</given-names></name> <name><surname>Izquierdo</surname><given-names>A.</given-names></name></person-group> (<year>2019</year>). <article-title>Adaptive learning under expected and unexpected uncertainty</article-title>. <source>Nat. Rev. Neurosci.</source> <volume>20</volume>, <fpage>635</fpage>&#x2013;<lpage>644</lpage>. doi: <pub-id pub-id-type="doi">10.1038/s41583-019-0180-y</pub-id>, PMID: <pub-id pub-id-type="pmid">31147631</pub-id></citation></ref>
<ref id="ref32"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Stone</surname><given-names>B.</given-names></name> <name><surname>Dowling</surname><given-names>S.</given-names></name> <name><surname>Cameron</surname><given-names>A.</given-names></name></person-group> (<year>2019</year>). <article-title>Cognitive impairment and homelessness: a scoping review</article-title>. <source>Health Soc. Care Commun.</source> <volume>27</volume>, <fpage>e125</fpage>&#x2013;<lpage>e142</lpage>. doi: <pub-id pub-id-type="doi">10.1111/hsc.12682</pub-id>, PMID: <pub-id pub-id-type="pmid">30421478</pub-id></citation></ref>
<ref id="ref33"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Subramanian</surname><given-names>A.</given-names></name> <name><surname>Chitlangia</surname><given-names>S.</given-names></name> <name><surname>Baths</surname><given-names>V.</given-names></name></person-group> (<year>2022</year>). <article-title>Reinforcement learning and its connections with neuroscience and psychology</article-title>. <source>Neural Netw.</source> <volume>145</volume>, <fpage>271</fpage>&#x2013;<lpage>287</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.neunet.2021.10.003</pub-id>, PMID: <pub-id pub-id-type="pmid">34781215</pub-id></citation></ref>
<ref id="ref34"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Thomas</surname><given-names>M. M. C.</given-names></name> <name><surname>Waldfogel</surname><given-names>J.</given-names></name></person-group> (<year>2022</year>). <article-title>What kind of &#x201C;poverty&#x201D; predicts CPS contact: income, material hardship, and differences among racialized groups</article-title>. <source>Child Youth Serv. Rev.</source> <volume>136</volume>:<fpage>106400</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.childyouth.2022.106400</pub-id>, PMID: <pub-id pub-id-type="pmid">35462724</pub-id></citation></ref>
<ref id="ref35"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Torpy</surname><given-names>J. M.</given-names></name> <name><surname>Lynm</surname><given-names>C.</given-names></name> <name><surname>Glass</surname><given-names>R. M.</given-names></name></person-group> (<year>2007</year>). <article-title>Poverty and health</article-title>. <source>JAMA</source> <volume>298</volume>:<fpage>1968</fpage>. doi: <pub-id pub-id-type="doi">10.1001/jama.298.16.1968</pub-id>, PMID: <pub-id pub-id-type="pmid">17954551</pub-id></citation></ref>
<ref id="ref36"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tottenham</surname><given-names>N.</given-names></name></person-group> (<year>2009</year>). <article-title>A review of adversity, the amygdala and the hippocampus: a consideration of developmental timing</article-title>. <source>Front. Hum. Neurosci.</source> <volume>3</volume>:<fpage>68</fpage>. doi: <pub-id pub-id-type="doi">10.3389/neuro.09.068.2009</pub-id>, PMID: <pub-id pub-id-type="pmid">20161700</pub-id></citation></ref>
<ref id="ref37"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname><given-names>J.</given-names></name> <name><surname>Li</surname><given-names>D.</given-names></name> <name><surname>Zhang</surname><given-names>W.</given-names></name></person-group> (<year>2010</year>). <article-title>Adolescents&#x2019; family financial difficulty and social adaptation: coping efficacy of compensatory, mediation, and moderation effects</article-title>. <source>J. Beijing Norm Univ. (Soc. Sci.)</source> <volume>4</volume>, <fpage>22</fpage>&#x2013;<lpage>32</lpage>.</citation></ref>
<ref id="ref38"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>White</surname><given-names>S. F.</given-names></name> <name><surname>Nusslock</surname><given-names>R.</given-names></name> <name><surname>Miller</surname><given-names>G. E.</given-names></name></person-group> (<year>2022</year>). <article-title>Low socioeconomic status is associated with a greater neural response to both rewards and losses</article-title>. <source>J. Cogn. Neurosci.</source> <volume>34</volume>, <fpage>1939</fpage>&#x2013;<lpage>1951</lpage>. doi: <pub-id pub-id-type="doi">10.1162/jocn_a_01821</pub-id>, PMID: <pub-id pub-id-type="pmid">35061015</pub-id></citation></ref>
<ref id="ref39"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xue</surname><given-names>G.</given-names></name> <name><surname>Xue</surname><given-names>F.</given-names></name> <name><surname>Droutman</surname><given-names>V.</given-names></name> <name><surname>Lu</surname><given-names>Z.-L.</given-names></name> <name><surname>Bechara</surname><given-names>A.</given-names></name> <name><surname>Read</surname><given-names>S.</given-names></name></person-group> (<year>2013</year>). <article-title>Common neural mechanisms underlying reversal learning by reward and punishment</article-title>. <source>PLoS One</source> <volume>8</volume>:<fpage>e82169</fpage>. doi: <pub-id pub-id-type="doi">10.1371/journal.pone.0082169</pub-id>, PMID: <pub-id pub-id-type="pmid">24349211</pub-id></citation></ref>
<ref id="ref40"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yacubian</surname><given-names>J.</given-names></name> <name><surname>Gl&#x00E4;scher</surname><given-names>J.</given-names></name> <name><surname>Schroeder</surname><given-names>K.</given-names></name> <name><surname>Sommer</surname><given-names>T.</given-names></name> <name><surname>Braus</surname><given-names>D. F.</given-names></name> <name><surname>B&#x00FC;chel</surname><given-names>C.</given-names></name></person-group> (<year>2006</year>). <article-title>Dissociable systems for gain- and loss-related value predictions and errors of prediction in the human brain</article-title>. <source>J. Neurosci.</source> <volume>26</volume>, <fpage>9530</fpage>&#x2013;<lpage>9537</lpage>. doi: <pub-id pub-id-type="doi">10.1523/JNEUROSCI.2915-06.2006</pub-id>, PMID: <pub-id pub-id-type="pmid">16971537</pub-id></citation></ref>
</ref-list>
</back>
</article>