<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Psychol.</journal-id>
<journal-title>Frontiers in Psychology</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Psychol.</abbrev-journal-title>
<issn pub-type="epub">1664-1078</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpsyg.2025.1522740</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Psychology</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>EuleApp&#x00A9;: a computerized adaptive assessment tool for early literacy skills</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name><surname>Yumus</surname> <given-names>Melike</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2611497/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Stuhr</surname> <given-names>Christina</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/929269/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Meindl</surname> <given-names>Marlene</given-names></name>
<xref ref-type="aff" rid="aff3"><sup>3</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2599018/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Leuschner</surname> <given-names>Haug</given-names></name>
<xref ref-type="aff" rid="aff4"><sup>4</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Jungmann</surname> <given-names>Tanja</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/2058751/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>Department of Special Needs Education and Rehabilitation, Carl von Ossietzky University of Oldenburg</institution>, <addr-line>Oldenburg</addr-line>, <country>Germany</country></aff>
<aff id="aff2"><sup>2</sup><institution>Faculty of Philosophy, Institute for Sports Science, University of Rostock</institution>, <addr-line>Rostock</addr-line>, <country>Germany</country></aff>
<aff id="aff3"><sup>3</sup><institution>Department of Special Education and Rehabilitation, University of Rostock</institution>, <addr-line>Rostock</addr-line>, <country>Germany</country></aff>
<aff id="aff4"><sup>4</sup><institution>DHL Data Science Seminare GmbH</institution>, <addr-line>K&#x00F6;ln</addr-line>, <country>Germany</country></aff>
<author-notes>
<fn fn-type="edited-by" id="fn0001">
<p>Edited by: Ioannis Tsaousis, National and Kapodistrian University of Athens, Greece</p>
</fn>
<fn fn-type="edited-by" id="fn0002">
<p>Reviewed by: Kelvin Fai Hong Lui, Lingnan University, Hong Kong SAR, China</p>
<p>Penelope Collins, University of California, Irvine, United States</p>
<p>Guher Gorgun, University of Kiel, Germany</p>
</fn>
<corresp id="c001">&#x002A;Correspondence: Melike Yumus, <email>melike.yumus@uni-oldenburg.de</email></corresp>
</author-notes>
<pub-date pub-type="epub">
<day>25</day>
<month>04</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>16</volume>
<elocation-id>1522740</elocation-id>
<history>
<date date-type="received">
<day>04</day>
<month>11</month>
<year>2024</year>
</date>
<date date-type="accepted">
<day>28</day>
<month>02</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2025 Yumus, Stuhr, Meindl, Leuschner and Jungmann.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Yumus, Stuhr, Meindl, Leuschner and Jungmann</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<sec>
<title>Introduction</title>
<p>Ample evidence indicates that assessing children&#x2019;s early literacy skills is crucial for later academic success. This assessment enables the provision of necessary support and materials while engaging them in the culture of print and books before school entry. However, relatively few assessment tools are available to identify early literacy skills, such as concepts of print, print awareness, phonological awareness, word awareness, alphabet knowledge, and early reading. The digital landscape presents new opportunities to enhance these assessments and provide enriching early literacy experiences. This study examines the psychometric properties of an adaptive assessment tool, EuLeApp&#x00A9;, focusing on its reliability and concurrent validity.</p>
</sec>
<sec>
<title>Methods</title>
<p>Data involved 307 German kindergarten children (M<sub>age</sub> = 64 months old, range = 45&#x2013;91). A Computerized Adaptive Testing (CAT) method, grounded in Item Response Theory (IRT), was employed to develop an adaptive digital tool for assessing early literacy competencies. We utilized an automatic item selection procedure based on item difficulty and discrimination parameters for the 183-item pool to ensure a precise and efficient assessment tailored to each child&#x2019;s ability level.</p>
</sec>
<sec>
<title>Results</title>
<p>The 4-parameter Logistic (4PL) model was identified as the best-fitting model for adaptive assessment, providing the highest precision in estimating children&#x2019;s abilities within this framework.</p>
</sec>
<sec>
<title>Discussions</title>
<p>The findings support the idea that the adaptive digital-based assessment tool EuLeApp&#x00A9; can be used to assess early literacy skills. It also provides a foundation for offering individualized and adaptable learning opportunities embedded in daily routines in daycare centers.</p>
</sec>
</abstract>
<kwd-group>
<kwd>early literacy</kwd>
<kwd>digital assessment</kwd>
<kwd>preschool age</kwd>
<kwd>item response theory</kwd>
<kwd>computerized adaptive test</kwd>
<kwd>psychometric validation</kwd>
</kwd-group>
<contract-num rid="cn1">01NV2105A</contract-num>
<contract-sponsor id="cn1">German Federal Ministry of Education and Research</contract-sponsor>
<counts>
<fig-count count="6"/>
<table-count count="7"/>
<equation-count count="0"/>
<ref-count count="96"/>
<page-count count="18"/>
<word-count count="13406"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Quantitative Psychology and Measurement</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec1">
<label>1</label>
<title>Introduction</title>
<p>Many studies have highlighted significant differences in reading and writing outcomes between strong and weak readers from early years through high school (<xref ref-type="bibr" rid="ref10">Buckingham et al., 2013</xref>; <xref ref-type="bibr" rid="ref51">McArthur et al., 2020</xref>). For instance, in Germany, recent investigations reveal that almost two in five fourth graders score &#x201C;below basic&#x201D; in their reading and writing skills, indicating they struggle to read and understand simple texts (<xref ref-type="bibr" rid="ref30">European Commission, 2023</xref>; <xref ref-type="bibr" rid="ref52">McElvany et al., 2023</xref>). This finding underscores a broader issue in educational systems worldwide: the failure to identify children at risk of reading difficulties early enough to provide timely and adequate interventions (<xref ref-type="bibr" rid="ref2">Adams, 1994</xref>; <xref ref-type="bibr" rid="ref17">Catts et al., 2001</xref>; <xref ref-type="bibr" rid="ref42">Justice and Ezell, 2001</xref>). This disparity also ties into the dyslexia paradox, where the most effective interventions occur before a child experiences academic failure (<xref ref-type="bibr" rid="ref73">Shaywitz, 1998</xref>). Addressing this issue requires reliable, efficient, and adaptive diagnostic tools capable of predicting early literacy skills from preschool onward. Such tools would enable educators to detect literacy difficulties early, providing the foundation for timely and targeted interventions (<xref ref-type="bibr" rid="ref14">Care et al., 2018</xref>; <xref ref-type="bibr" rid="ref33">Gonski et al., 2018</xref>). This approach also aligns with ecological and sociocultural approaches to assessment, which emphasize a child&#x2019;s &#x201C;readiness to learn&#x201D; and the zone of proximal development, where scaffolding by teachers supports learning from the familiar to the unfamiliar (<xref ref-type="bibr" rid="ref001">Vygotsky, 1978</xref>). Through these approaches, assessment becomes a means to monitor a child&#x2019;s progress over time, fostering an adaptive learning environment. The app&#x2019;s computerized adaptive testing (CAT) framework dynamically adjusts to each child&#x2019;s ability level, making it possible to identify literacy difficulties, including dyslexia, early on. This approach is critical for reducing the gap between strong and weak readers by tailoring assessments and providing opportunities for interventions to each child&#x2019;s unique needs, thus preventing the long-term consequences of untreated literacy deficits.</p>
<p>The accurate and timely assessment of early literacy skills is essential for early identification of learning gaps, customizing daily literacy activities, monitoring progress, and making data-driven adjustments to narrow achievement disparities. Thus, this project aimed to develop an adaptive assessment tool and to demonstrate the reliability and concurrent validity of EuLeApp&#x00A9; in a German sample of children aged 4&#x2013;7 years. The adaptive approach is essential because it allows for precise performance estimates by adjusting the difficulty and number of items based on each child&#x2019;s ability level (<xref ref-type="bibr" rid="ref81">Weiss, 2004</xref>). By leveraging CAT, EuLeApp&#x00A9; maximizes both efficiency and accuracy, making it a valuable tool for early literacy assessment. Technology-based solutions, such as the EuLeApp&#x00A9;, hold great potential in this area.</p>
<sec id="sec2">
<label>1.1</label>
<title>The importance of individualized assessment of early literacy skills</title>
<p>Early literacy encompasses a range of skills related to oral language, phonological awareness, word awareness, print awareness, concepts of print, alphabet knowledge, and narrative abilities that develop before children formally learn to read and write. Phonological awareness refers to the ability to detect the smallest sound units within words. When children develop this skill, they understand that language can be analyzed and manipulated, a key milestone in their literacy journey toward understanding the concept of words (<xref ref-type="bibr" rid="ref31">Gee, 2012</xref>; <xref ref-type="bibr" rid="ref76">Snow and Dickinson, 1991</xref>). Word awareness involves recognizing that words, as elements of language, have properties independent of their meaning. For example, children learn to connect printed words in their oral vocabulary while learning to read. This awareness also includes understanding word boundaries and what constitutes a word (<xref ref-type="bibr" rid="ref42">Justice and Ezell, 2001</xref>). The term print awareness refers to the knowledge that print carries meaning, and differs structurally from other sign systems (e.g., numbers). To develop print awareness, it is crucial for young children to be exposed to letters and written text in their environment. Familiarity with books and print culture also involves understanding the characteristics of books and how they are read, which relates to concepts of print (<xref ref-type="bibr" rid="ref55">Meindl and Jungmann, 2019</xref>; <xref ref-type="bibr" rid="ref66">Piasta et al., 2012</xref>). Alphabet knowledge entails recognizing the characteristics of different graphemes and associating them with their corresponding phonemes. Early reading skills include understanding that words are made up of graphemes (<xref ref-type="bibr" rid="ref28">Elimelech and Aram, 2020</xref>; <xref ref-type="bibr" rid="ref56">Morrow, 2007</xref>). Narrative skills reflect children&#x2019;s ability to produce a fictional or real account of a temporally sequenced experience or event (<xref ref-type="bibr" rid="ref29">Engel, 1995</xref>). All these aspects represent crucial milestones in the successful development of reading and writing (<xref ref-type="bibr" rid="ref42">Justice and Ezell, 2001</xref>).</p>
<p>Early literacy tasks such as phonological awareness, print awareness, and word awareness reflect how children process and integrate the relationships between various linguistic rule systems on the metacognitive level. <xref ref-type="bibr" rid="ref57">Nagy and Anderson (1995)</xref> documented that metalinguistic awareness helps young children to become aware of the structure of their writing system and its relationship to their spoken language. For instance, children with higher metalinguistic awareness perform better on tasks related to concepts of print and phonological awareness, both of which are strong predictors of reading success (<xref ref-type="bibr" rid="ref21">Chaney, 1994</xref>). Without proper assessments, children in need of additional support may be overlooked, leading to long-term consequences such as ongoing academic difficulties, low self-esteem, and limited future opportunities (<xref ref-type="bibr" rid="ref8">Brassard and Boehm, 2007</xref>; <xref ref-type="bibr" rid="ref43">Justice et al., 2002</xref>; <xref ref-type="bibr" rid="ref62">OECD, 2013</xref>). Screening early literacy skills with standardized assessments has gained significant attention in early childhood for various reasons: (a) to evaluate a child&#x2019;s strengths and weaknesses in specific areas, (b) to identify key target skills and provide tailored support, (c) to structure educational programs, (d) to monitor children&#x2019;s progress over time, and (e) to improve educational outcomes by facilitating a smoother transition to school. Assessment tools also offer significant benefits at different stages of assessment, from testing to linking appropriate interventions, by collaborating with children&#x2019;s parents and other stakeholders to help overcome disadvantages. Despite the benefits of standardized testing, research shows that data-driven decision-making remains significantly underutilized in early education. Children enter kindergarten with a wide range of literacy and language abilities (<xref ref-type="bibr" rid="ref16">Catts et al., 2002</xref>), making individualized feedback essential (<xref ref-type="bibr" rid="ref8">Brassard and Boehm, 2007</xref>; <xref ref-type="bibr" rid="ref53">McLachlan et al., 2018</xref>; <xref ref-type="bibr" rid="ref77">Snow and van Hemel, 2008</xref>). To address these needs, computerized adaptive testing (CAT) provides an innovative solution for assessing children&#x2019;s literacy skills. By dynamically adjusting to each child&#x2019;s ability level, CAT improves measurement precision while reducing the required test items. As a cutting-edge tool, EuLeApp&#x00A9; leverages CAT to streamline the assessment process, ensuring both accuracy and efficiency. This adaptability allows practitioners to deliver individualized feedback, enabling more targeted interventions and improving educational outcomes.</p>
</sec>
<sec id="sec3">
<label>1.2</label>
<title>The importance of innovative assessment tools</title>
<p>Given the ongoing expansion of digital media in educational settings (<xref ref-type="bibr" rid="ref63">OECD, 2015</xref>) and the demand for more effective assessment methods, researchers have increasingly focused on the potential of digital tools for enhancing educational processes. These tools offer greater efficiency (<xref ref-type="bibr" rid="ref49">Marsh et al., 2020</xref>) and provide visually engaging reports for tracking learning progress (<xref ref-type="bibr" rid="ref60">Neumann et al., 2019</xref>). The use of digital assessment tools is increasing for numerous reasons, such as technological portability and ease, the touchscreen&#x2019;s tremendous potential to reach young children, and the need for mobile, adaptable, and accessible assessment tools (<xref ref-type="bibr" rid="ref60">Neumann et al., 2019</xref>). Furthermore, online assessment tools have become widespread and are actively used. For example, <xref ref-type="bibr" rid="ref40">Ho et al. (2024)</xref> developed a short-term online test to assess word reading. This test, which covers a wide age range, is primarily designed for individuals who can read to some degree at a basic level (ages 7 and above). Tablet-based assessments, on the other hand, allow for greater control over the evaluation process for younger children and may be a more suitable alternative. Additionally, app-based assessments offer the advantage of functioning offline, making them more accessible in diverse settings such as homes, preschools, and clinics, where stable internet access is not always available.</p>
<p>App-based assessments can be administered anywhere with a suitable device, making it easier for educators and researchers to assess children in various settings, including homes, classrooms, and clinics. Another benefit of mobile media devices is that they can help advance the goal of reaching many children for educational opportunities and equity because of their low costs and good accessibility (<xref ref-type="bibr" rid="ref39">Hirsh-Pasek et al., 2015</xref>). Additionally, traditional paper-pencil assessments usually require considerable time, effort, and expertise, such as organizing, rewriting, and preparing children for the test in person (<xref ref-type="bibr" rid="ref70">Schildkamp and Kuiper, 2010</xref>). Also, some assessment procedures rely on contextual factors such as observation, making it challenging to maintain strict objectivity (<xref ref-type="bibr" rid="ref38">Hindman et al., 2020</xref>). Moreover, children may not be able to show their best performance when assessed by someone they do not know (<xref ref-type="bibr" rid="ref36">Halliday et al., 2018</xref>). In contrast, computerized assessments offer the potential to gain insights into children&#x2019;s responses, such as disengagement, rapid guessing, or unexpected answers (<xref ref-type="bibr" rid="ref11">Bulut and Cormier, 2018</xref>; <xref ref-type="bibr" rid="ref46">Lee and Jia, 2014</xref>; <xref ref-type="bibr" rid="ref84">Wise and Kong, 2005</xref>). Thus, app-based assessments enable teachers to receive prompt feedback, facilitating more effective support and remediation. However, technology-based assessments also raise concerns regarding developmental appropriateness, item development, psychometric validity, and teacher training (<xref ref-type="bibr" rid="ref11">Bulut and Cormier, 2018</xref>; <xref ref-type="bibr" rid="ref60">Neumann et al., 2019</xref>). Tablet-based assessments are common in many schools across the United States, where they are used to assess and monitor students&#x2019; performance in mathematics, reading, and science throughout the school year (<xref ref-type="bibr" rid="ref23">Davey, 2005</xref>; <xref ref-type="bibr" rid="ref72">Sharkey and Murnane, 2006</xref>). Despite the increasing integration of technology in education, the development of digital tools specifically for early literacy assessment remains limited, underscoring the need for further innovation. Emerging research shows that only a few app-based tools are available to assess children&#x2019;s language and literacy skills in the early years. One such tool is Logometro<sup>&#x00AE;</sup>, a reliable app-based test that evaluates children&#x2019;s phonological awareness, listening comprehension, vocabulary, narrative skills, speech, morphological awareness, and pragmatic skills (<xref ref-type="bibr" rid="ref3">Antoniou et al., 2022</xref>). Administered through a specially developed Android app, Logometro<sup>&#x00AE;</sup> allows for accurate directional vocalization and easy capture of children&#x2019;s responses via touchscreens and direct recordings (<xref ref-type="bibr" rid="ref3">Antoniou et al., 2022</xref>). Another innovative app, QUILS (<xref ref-type="bibr" rid="ref32">Golinkoff et al., 2017</xref>) focuses on assessing children&#x2019;s language-learning processes, offering insights into how they acquire new words and grammatical structures. <xref ref-type="bibr" rid="ref32">Golinkoff et al. (2017)</xref> demonstrated that a technology-based assessment tool measuring children&#x2019;s phonological awareness and letter knowledge was efficient in terms of time usage and effectively differentiated these skills. One app was also developed by <xref ref-type="bibr" rid="ref59">Neumann (2018)</xref> to assess children&#x2019;s letter knowledge and vocabulary using both expressive and receptive response formats. The valuable insights gained from such innovative assessment tools can make them more attractive, accurate, and accessible for educators and parents.</p>
</sec>
<sec id="sec4">
<label>1.3</label>
<title>Computerized adaptive testing (CAT)</title>
<p>Computerized Adaptive Tests (CAT) dynamically adjust to a child&#x2019;s ability by selecting test items based on their responses. Unlike traditional fixed-form tests, CAT tailors the difficulty of each item to match the child&#x2019;s performance, selecting more challenging questions after correct answers and easier ones after incorrect responses (<xref ref-type="bibr" rid="ref54">Meijer and Nering, 1999</xref>; <xref ref-type="bibr" rid="ref82">Weiss and Kingsbury, 1984</xref>). This individualized approach helps maintain engagement and ensure a more accurate assessment of abilities (<xref ref-type="bibr" rid="ref79">Tomasik et al., 2018</xref>; <xref ref-type="bibr" rid="ref83">Wise, 2014</xref>). CAT relies on an Item Response Theory (IRT) calibrated item bank, selecting items sequentially to estimate a child&#x2019;s ability (<italic>&#x03B8;</italic>) more precisely with fewer items than conventional tests (<xref ref-type="bibr" rid="ref19">Chalmers, 2015</xref>; <xref ref-type="bibr" rid="ref45">Keuning and Verhoeven, 2008</xref>). The efficiency of CAT is enhanced by large item banks, allowing the test to adapt to individual performance effectively (<xref ref-type="bibr" rid="ref58">Nelson et al., 2017</xref>). IRT provides a robust framework for interpreting test scores, as it allows for accurate item selection and predicts the likelihood that a child will respond correctly based on their skill level (<xref ref-type="bibr" rid="ref6">Bjorner et al., 2007</xref>; <xref ref-type="bibr" rid="ref4">Baker, 2001</xref>). One key advantage of IRT is that item parameters remain stable across different samples, meaning they are not dependent on the specific group being tested (<xref ref-type="bibr" rid="ref47">Magis and Barrada, 2017</xref>). Additionally, IRT offers a precision measure or standard error for each skill estimate (<xref ref-type="bibr" rid="ref37">He and Min, 2017</xref>), providing insight into the accuracy of the assessment across varying levels of ability (<xref ref-type="bibr" rid="ref80">Weiss, 1982</xref>).</p>
</sec>
<sec id="sec5">
<label>1.4</label>
<title>Present study</title>
<p>The quality of assessment tools is important for educators to understand what children are expected to learn before formal school readiness. A primary aim of the current study is to adapt a digital assessment tool from the paper-pencil-based EuLe 4&#x2013;5 assessment (<xref ref-type="bibr" rid="ref55">Meindl and Jungmann, 2019</xref>), a standardized tool designed to assess narrative and early literacy skills in German children aged 4;0 to 5;11&#x202F;years. Based on this goal, the study also aims to validate the EuLeApp&#x00A9; as a digital, adaptive assessment tool for children aged between 4;0 and 7;11&#x202F;years. For this, a Multidimensional Computerized Adaptive Test (MCAT) was used based on Item Response Theory (IRT), allowing for individualized and precise measurement of children&#x2019;s early literacy skills across multiple dimensions. An item pool was constructed through calibration based on the content of the items, and model fit was estimated and established using an IRT model.</p>
<p>These research questions will be addressed as follows:<list list-type="order">
<list-item>
<p>Does the EuLeApp&#x00A9; screening tool accurately assess early literacy skills in children aged 4 to 7&#x202F;years?</p>
</list-item>
<list-item>
<p>How can item response theory (IRT) be used to optimize item difficulty in computerized adaptive testing (CAT) for assessing children&#x2019;s early literacy skills?</p>
</list-item>
</list></p>
</sec>
</sec>
<sec sec-type="materials|methods" id="sec6">
<label>2</label>
<title>Materials and method</title>
<sec id="sec7">
<label>2.1</label>
<title>Sample</title>
<p>The sample consisted of <italic>N</italic>&#x202F;=&#x202F;307 kindergarten children (M<sub>age</sub>&#x202F;=&#x202F;64&#x202F;months, range&#x202F;=&#x202F;45&#x2013;91) before entering formal schooling in Mecklenburg-West Pomerania and Lower Saxony in Germany. The sample distribution of boys (<italic>n</italic>&#x202F;=&#x202F;170, 55.4%) and girls (<italic>n</italic>&#x202F;=&#x202F;137, 44.6%) was approximately equal. In terms of the distribution of children&#x2019;s ages, 31.3% were 4 years old, 47.2% were 5 years old, 13.0% were 6 years old, and 8.5% were 7 years old. Data were primarily collected from kindergartens in middle- and high-socioeconomic regions. All parents were informed about the study, and written consent was obtained. <xref ref-type="table" rid="tab1">Table 1</xref> provides an overview of the participant demographics, including age distribution, gender, and other key characteristics.</p>
<table-wrap position="float" id="tab1">
<label>Table 1</label>
<caption>
<p>Participant demographics (<italic>N</italic>&#x202F;=&#x202F;307).</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Variable</th>
<th align="left" valign="top">Category</th>
<th align="center" valign="top">
<italic>N</italic>
</th>
<th align="center" valign="top">Frequency (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">State</td>
<td align="left" valign="middle">Mecklenburg-West Pomerania</td>
<td align="center" valign="middle">171</td>
<td align="center" valign="middle">55.7</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">Lower Saxony</td>
<td align="center" valign="middle">136</td>
<td align="center" valign="middle">44.3</td>
</tr>
<tr>
<td align="left" valign="middle">Gender</td>
<td align="left" valign="middle">Boys</td>
<td align="center" valign="middle">170</td>
<td align="center" valign="middle">55.4</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">Girls</td>
<td align="center" valign="middle">137</td>
<td align="center" valign="middle">44.6</td>
</tr>
<tr>
<td align="left" valign="middle">Age group</td>
<td align="left" valign="middle">4&#x202F;years</td>
<td align="center" valign="middle">96</td>
<td align="center" valign="middle">31.3</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">5&#x202F;years</td>
<td align="center" valign="middle">145</td>
<td align="center" valign="middle">47.2</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">6&#x202F;years</td>
<td align="center" valign="middle">40</td>
<td align="center" valign="middle">13.0</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">7&#x202F;years</td>
<td align="center" valign="middle">26</td>
<td align="center" valign="middle">8.5</td>
</tr>
<tr>
<td align="left" valign="middle">Parental education</td>
<td align="left" valign="middle">Secondary education</td>
<td align="center" valign="middle">123</td>
<td align="center" valign="middle">40.1</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">Vocational training</td>
<td align="center" valign="middle">53</td>
<td align="center" valign="middle">17.3</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">Higher education</td>
<td align="center" valign="middle">131</td>
<td align="center" valign="middle">42.6</td>
</tr>
<tr>
<td align="left" valign="middle">Linguistic diversity</td>
<td align="left" valign="middle">Monolingual (German)</td>
<td align="center" valign="middle">253</td>
<td align="center" valign="middle">82.4</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">Bilingual</td>
<td align="center" valign="middle">24</td>
<td align="center" valign="middle">7.8</td>
</tr>
<tr>
<td/>
<td align="left" valign="middle">Missing data</td>
<td align="center" valign="middle">30</td>
<td align="center" valign="middle">9.8</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="sec8">
<label>2.2</label>
<title>Measures</title>
<sec id="sec9">
<label>2.2.1</label>
<title>Early literacy assessment app (EuLeApp&#x00A9;)</title>
<p>To assess early literacy between the ages of 4 and 7 years, we administered the EuLeApp&#x00A9;, a digital multiple-choice test developed to measure key early literacy content areas, including the following: (a) the concepts of print, (b) print awareness, (c) word awareness, (d) phonological awareness, (e) alphabet knowledge, and (f) early reading (g) narrative skills (<xref ref-type="fig" rid="fig1">Figure 1</xref>).</p>
<fig position="float" id="fig1">
<label>Figure 1</label>
<caption>
<p>EuLeApp&#x00A9; subtests with sample test items.</p>
</caption>
<graphic xlink:href="fpsyg-16-1522740-g001.tif"/>
</fig>
</sec>
<sec id="sec10">
<label>2.2.2</label>
<title>EuLeApp&#x00A9; features</title>
<p>During the development of the EuLeApp&#x00A9; prototype, four key design features were incorporated to support the self-assessment process: definition, modeling, practice, and motivation (<xref ref-type="bibr" rid="ref15">Carson et al., 2015</xref>). The &#x201C;definition&#x201D; feature provides a brief explanation of what the child will do and how the task works. While the &#x201C;modeling&#x201D; feature demonstrates how to use the system and what is expected of the child (e.g., how to interact with the tool and respond to questions), the &#x201C;practice&#x201D; feature allows the child to practice how the system works before the actual assessment. The goal of the &#x201C;motivation&#x201D; feature is to motivate and encourage children to keep them engaged in completing the test.</p>
<p>This digital assessment tool can be completed in a single session (&#x223C;20&#x202F;min). It requires no formal training and can be automatically administered and scored using its software program, except for the picture story for the narrative part &#x201C;Seagull Marius.&#x201D; For only this narrative part, the pedagogical professionals have to analyze the realized macrostructure of the child&#x2019;s narration with the help of a protocol sheet.</p>
<p>The EuLeApp&#x00A9;&#x2019;s subscales, each tapping into different and overlapping skills, were developed to provide teachers, pedagogic staff, researchers, and parents with sufficient information to interpret children&#x2019;s early literacy performance. The assessment follows a multiple-choice format, which is widely recognized as an effective method for direct measurement (<xref ref-type="bibr" rid="ref35">Haladyna, 2013</xref>; <xref ref-type="bibr" rid="ref67">Raymond et al., 2019</xref>). Accordingly, each child is shown with a picture or figure with four response options on the screen (one target response, three distractors). Within these six subscales (<xref ref-type="fig" rid="fig1">Figure 1</xref>), we presented children with 175 items in total. Item indicators are equivalent to items in a conventional test. Item selection considered linguistic (e.g., bilingual, multilingual), socioeconomic, and cultural diversity (<xref ref-type="bibr" rid="ref002">Levine et al., 2020</xref>). In this context, we used short and clear instructions. Besides being more effective than open-ended questions, another important reason for structuring the test in a multiple-choice format is considering equity for children from different backgrounds (<xref ref-type="bibr" rid="ref003">Bruder, 1993</xref>; <xref ref-type="bibr" rid="ref60">Neumann et al., 2019</xref>).</p>
<p><italic>Concepts of print</italic>: This task comprises 40 items, and children need to identify print-related tasks such as reading directions from left to right or where to start reading (Cronbach&#x2019;s <italic>&#x03B1;</italic>&#x202F;=&#x202F;0.93). Based on the questions, children are required to tap on the correct part of the screen.</p>
<p><italic>Print awareness</italic>: The assessment of print awareness comprises a 19-item subtest (Cronbach&#x2019;s <italic>&#x03B1;</italic>&#x202F;=&#x202F;0.76). The task of the children is to distinguish between words, icons, and symbols (&#x201C;Tap on the letter&#x201D;; &#x201C;Tap on writing&#x201D;). Each child is presented with four pictures, one of which is the target and the other three serving as distractors, and is expected to tap either on the word or the letter corresponding to the target.</p>
<p><italic>Words awareness</italic>: This task consists of 11 items that increase in difficulty gradually (Cronbach&#x2019;s &#x03B1;&#x202F;=&#x202F;0.79). During this assessment, children encounter short texts and are directed to tap on specific elements, such as the first, second, or space between two words (&#x201C;Tap on the space between two words&#x201D;).</p>
<p><italic>Phonological awareness</italic>: The phonological awareness subtest comprises 29 items (Cronbach&#x2019;s &#x03B1;&#x202F;=&#x202F;0.72) and assesses the ability of two key components: synthesizing syllables and phonemes to form words and analyzing words&#x2019; syllabic and phonemic structure. For example, tasks include identifying the initial sounds of words (e.g., &#x201C;Here you can see three pictures: Grandma, Mum, Apple. Touch the picture that starts with /m/&#x201D;).</p>
<p><italic>Alphabet knowledge</italic>: This subtest consists of 45 items (Cronbach&#x2019;s &#x03B1;&#x202F;=&#x202F;0.89). To evaluate children&#x2019;s alphabet knowledge, they are presented with phonetic realizations of letters and are asked to select the corresponding letter from a set of four options (e.g., &#x201C;Tap the /m/&#x201D;).</p>
<p><italic>Early reading</italic>: In the early reading segment, children are asked to name the 36 letters of the alphabet (Cronbach&#x2019;s &#x03B1;&#x202F;=&#x202F;0.96). Items such as &#x201C;Tap on /am/,&#x201D; &#x201C;Tap on /mama/&#x201D; (receptive segment), and &#x201C;What is written here?&#x201D; (productive segment) assess children&#x2019;s first receptive and productive reading abilities on the syllable and word level.</p>
<p><xref ref-type="table" rid="tab2">Table 2</xref> provides descriptive statistics, reliability estimates, and distribution metrics (skewness and kurtosis) for each subscale. While Cronbach&#x2019;s &#x03B1; values indicate good internal consistency across the EuLeApp&#x00A9; scales, particularly for Concepts of Print, Alphabet Knowledge, and Early Reading, the inter-item correlation values were comparatively lower. Additionally, skewness and kurtosis values indicate that most subscales exhibit approximately normal distributions, with no extreme deviations from normality, supporting the reliability and usability of the scales.</p>
<table-wrap position="float" id="tab2">
<label>Table 2</label>
<caption>
<p>Descriptive statistics and internal consistency of the scores in each subscale (<italic>N</italic>&#x202F;=&#x202F;307).</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Scale</th>
<th align="center" valign="top">
<italic>M</italic>
</th>
<th align="center" valign="top">
<italic>SD</italic>
</th>
<th align="center" valign="top"><italic>M</italic> IIC</th>
<th align="center" valign="top"><italic>SD</italic> IIC</th>
<th align="center" valign="top">Cronbach&#x2019;s <inline-formula>
<mml:math id="M1">
<mml:mi>&#x03B1;</mml:mi>
</mml:math>
</inline-formula></th>
<th align="center" valign="top">Skewness</th>
<th align="center" valign="top">Kurtosis</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">Concept of prints</td>
<td align="center" valign="middle">23.7</td>
<td align="center" valign="middle">10.4</td>
<td align="center" valign="middle">0.277</td>
<td align="center" valign="middle">0.156</td>
<td align="center" valign="middle">0.939</td>
<td align="center" valign="top">&#x2212;0.065</td>
<td align="center" valign="top">&#x2212;0.891</td>
</tr>
<tr>
<td align="left" valign="middle">Print awareness</td>
<td align="center" valign="middle">11.8</td>
<td align="center" valign="middle">3.7</td>
<td align="center" valign="middle">0.157</td>
<td align="center" valign="middle">0.117</td>
<td align="center" valign="middle">0.769</td>
<td align="center" valign="top">&#x2212;0.437</td>
<td align="center" valign="top">&#x2212;0.058</td>
</tr>
<tr>
<td align="left" valign="middle">Word awareness</td>
<td align="center" valign="middle">5.3</td>
<td align="center" valign="middle">3.2</td>
<td align="center" valign="middle">0.244</td>
<td align="center" valign="middle">0.106</td>
<td align="center" valign="middle">0.796</td>
<td align="center" valign="top">0.294</td>
<td align="center" valign="top">&#x2212;0.678</td>
</tr>
<tr>
<td align="left" valign="middle">Phon. Awareness</td>
<td align="center" valign="middle">18.9</td>
<td align="center" valign="middle">4.3</td>
<td align="center" valign="middle">0.082</td>
<td align="center" valign="middle">0.083</td>
<td align="center" valign="middle">0.729</td>
<td align="center" valign="top">0.821</td>
<td align="center" valign="top">0.373</td>
</tr>
<tr>
<td align="left" valign="middle">Alphabet knowledge</td>
<td align="center" valign="middle">19.3</td>
<td align="center" valign="middle">9.0</td>
<td align="center" valign="middle">0.152</td>
<td align="center" valign="middle">0.083</td>
<td align="center" valign="middle">0.892</td>
<td align="center" valign="top">0.782</td>
<td align="center" valign="top">0.274</td>
</tr>
<tr>
<td align="left" valign="middle">Early reading</td>
<td align="center" valign="middle">6.9</td>
<td align="center" valign="middle">8.9</td>
<td align="center" valign="middle">0.427</td>
<td align="center" valign="middle">0.113</td>
<td align="center" valign="middle">0.962</td>
<td align="center" valign="top">0.14</td>
<td align="center" valign="top">&#x2212;0.372</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>
<italic>M IIC, Mean of Inter-Item Correlation; SD IIC, Standard Deviation of Inter-Item Correlation.</italic>
</p>
</table-wrap-foot>
</table-wrap>
<p>In addition, the EuLeApp&#x00A9; includes a narrative measure in which children retell the story of Seagull Marius after viewing seven pictures. Therefore, this section of the assessment follows a different methodological approach, and its reliability evaluation is still ongoing.</p>
<p>The EuLeApp&#x00A9; was designed with the understanding that children can comfortably engage with tablets. The test begins with a short, child-friendly explanation and an example of how they will conduct it, and it provides a short practice. After an item and its audio are presented, children respond to the question by touching one of the options on the screen. The test then continues with children viewing items one by one. The software is configured so that children can respond flexibly.</p>
<p>The test automatically moves through subtests, and children view short, fun, animated scenes throughout the assessment as a break from the tasks. According to <xref ref-type="bibr" rid="ref13">Campbell and Jane (2012)</xref>, motivation is considered one way to stimulate children&#x2019;s engagement. Therefore, children are regularly given positive feedback (e.g., &#x201C;Well done!&#x201D; or &#x201C;Excellent!&#x201D;) and visual, colorful gift boxes after each subtest.</p>
<sec id="sec11">
<label>2.2.2.1</label>
<title>Scoring and coding</title>
<p>To describe the child&#x2019;s strengths and weaknesses, a specialized database architecture was used to record the child&#x2019;s interactions with each task on the tablet. An automated scoring system provides real-time feedback by evaluating whether each response is correct or incorrect. Most indicators are binary, designed to detect the presence or absence of a correct answer for each item. By automating this process, the system can quickly and accurately generate feedback highlighting the child&#x2019;s strengths and identifying areas needing improvement. In other words, each early literacy domain is scored across a different number of items, and assessments of the children&#x2019;s answers, false (0) or true (1) scores, are determined by their success or failure on the task. The coded items are considered the primary data source for the scoring process. Evaluation indicators are classified based on the child&#x2019;s answers for each scale. Each assessment includes the child&#x2019;s total response time, item counts, and each child is coded with a unique ID code. To capture the required data, once the indicators are determined for items in every subscale, difficulty differences based on these item indicators are (automatically) determined.</p>
<p>Children&#x2019;s performance is reported as a score for each scale, allowing researchers and practitioners to monitor children&#x2019;s progress. Furthermore, to facilitate the classification of early literacy skill levels, the assessment results in a color-coded ranking list with &#x201C;traffic light&#x201D; analogy (<xref ref-type="bibr" rid="ref78">Templin and Henson, 2010</xref>). Children are categorized as &#x201C;red&#x201D; (at risk), &#x201C;orange&#x201D; (monitoring needed), or &#x201C;green&#x201D; (on track) based on their performance in key components of early literacy. This feature provides a viable way to identify children needing additional support, reinforcing the tool&#x2019;s ability to differentiate between skill levels and guide targeted interventions. The reports are also displayed on the children&#x2019;s profile pages and include short demographic information such as the child&#x2019;s age, gender, and kindergarten/school. The App&#x2019;s user-friendly interface allows for straightforward administration and scoring, making it a valuable resource for educators and researchers in early childhood education.</p>
</sec>
</sec>
<sec id="sec12">
<label>2.2.3</label>
<title>Language competence</title>
<p>Children&#x2019;s language competence was assessed using standardized German language tests, selected based on age: (a) &#x201C;Language level test for children aged 3&#x2013;5&#x202F;years&#x201D; (Sprachentwicklungstest f&#x00FC;r Kinder 3&#x2013;5 [SET 3&#x2013;5]; <xref ref-type="bibr" rid="ref65">Petermann et al., 2016</xref>), (b) &#x201C;Language level test for children aged 5&#x2013;10&#x202F;years&#x201D; (Sprachentwicklungstest f&#x00FC;r Kinder 5&#x2013;10 [SET 5&#x2013;10]; <xref ref-type="bibr" rid="ref64">Petermann, 2018</xref>). The SET 3&#x2013;5 consists of 12 subtests that measure a child&#x2019;s receptive language processing skills (understanding, recording), productive language processing skills (own speech acts), and auditory memory skills (language memory). The internal consistency for administered subtests from the SET 3&#x2013;5 ranged between <italic>&#x03B1;</italic>&#x202F;=&#x202F;0.70 and &#x03B1;&#x202F;=&#x202F;0.93 (<xref ref-type="bibr" rid="ref65">Petermann et al., 2016</xref>). The SET 5&#x2013;10 consists of 8 subtests to measure a child&#x2019;s vocabulary, semantic relations, processing speed, language comprehension, language production, grammar/ morphology, and auditory memory. The internal consistency for administered subtests ranges between &#x03B1;&#x202F;=&#x202F;0.71 and &#x03B1;&#x202F;=&#x202F;0.91 for the SET 5&#x2013;10 (<xref ref-type="bibr" rid="ref64">Petermann, 2018</xref>).</p>
</sec>
</sec>
<sec id="sec13">
<label>2.3</label>
<title>Procedures</title>
<p>Prior to the start of the study, university ethics approval was received, and permission to assess children was obtained from the head educators of a total of 15 kindergartens. Before administering the EuLeApp&#x00A9;, children&#x2019;s language skills were evaluated to ensure that they possessed sufficient language comprehension. This step was essential in preventing the misinterpretation of literacy test results due to underlying language receptive problems and ensuring that literacy performance was accurately measured. Then, we assessed early literacy skills in the daycare centers using the prototype of the EuLeApp&#x00A9; on a tablet. The test practitioners were master&#x2019;s students, PhD candidates, and postdoctoral fellows, all of whom completed two training sessions: one on understanding the assessment tool and its usage, and another on practical test implementation. All assessments were conducted individually. Before the test began, practitioners informed the children about the goal of the test and how it would help them, reassuring them that the test would not show everything they knew and could do and that they had plenty of time to answer the questions. Sitting beside the child, the test practitioners asked the child to practice tapping on the screen before beginning the assessment. Once the evaluation began, the practitioners did not answer the children&#x2019;s questions or provide any tips to ensure standardized administration. During the EuleApp&#x00A9; assessment procedure, standard administration and scoring procedures were followed.</p>
</sec>
<sec id="sec14">
<label>2.4</label>
<title>Data analysis strategy</title>
<p><xref ref-type="fig" rid="fig2">Figure 2</xref> outlines the structured process used for the CAT analysis in EuLeApp&#x00A9;: (a) Developing a calibrated item bank: Relevant items from the EuLe 4&#x2013;5 paper-based assessment tool were selected, categorized, and visualized to ensure consistency between the paper and digital formats. (b) Selection of starting items: A prototype was developed, and data were collected for item calibration, with model fit tested for accuracy. (c) Continuous estimation of a child&#x2019;s ability: A child&#x2019;s ability was continuously estimated during CAT simulation studies, applying a stop rule based on predefined precision criteria. (d) A final item pool was established based on simulations, integrated into EuLeApp&#x00A9;, and validated through reliable retest processes.</p>
<fig position="float" id="fig2">
<label>Figure 2</label>
<caption>
<p>The development of an adaptive app-based early literacy assessment tool.</p>
</caption>
<graphic xlink:href="fpsyg-16-1522740-g002.tif"/>
</fig>
<p>The R package mirtCAT was used for psychometric development of the app&#x2019;s multidimensional computerized adaptive test (MCAT) based on Item Response Theory (IRT). Initially, items were adapted from the EuLe 4&#x2013;5 test, whose content and construct validity had been previously established (<xref ref-type="bibr" rid="ref55">Meindl and Jungmann, 2019</xref>). A digital platform was then created to deliver the assessment, incorporating multimedia features (vocal instructions, touch interactions, and graphics) to enhance children&#x2019;s engagement and navigation. In the next step, data collection marked the first calibration phase of item pool development. A series of confirmatory IRT models were used to estimate item parameters and assess the effectiveness of the MCAT in providing individualized early literacy assessments. Both exploratory and confirmatory Item Factor Analysis (IFA) were conducted to validate the item structure, removing items misaligned with the identified factors to improve accuracy and reliability. Fit indices, including the Comparative Fit Index (CFI), Tucker&#x2013;Lewis Index (TLI), Root Mean Square Error of Approximation (RMSEA), and Standardized Root Mean Square Residual (SRMR), were applied to evaluate the calibration data against the proposed six-factor model (<xref ref-type="bibr" rid="ref9">Brown, 2015</xref>).</p>
<p>To address potential estimation challenges and improve model convergence, the model was divided to reduce the number of estimated parameters: the first submodel included scales 1&#x2013;4, while the second submodel included scales 5&#x2013;6. Using these submodels, exploratory and confirmatory IFA models were developed. Structural analysis indicated that items 9&#x2013;40 exhibited a bifactor structure (<xref ref-type="bibr" rid="ref22">Chen et al., 2006</xref>; <xref ref-type="bibr" rid="ref25">Dunn and McCray, 2020</xref>). This meant that the &#x201C;Concepts of Print&#x201D; scale could only be derived if the model incorporated two additional dimensions related to the item-specific use of numbers and images, which did not correlate with other dimensions. These were defined as &#x201C;Numerical Writing Awareness&#x201D; and &#x201C;Iconic Writing Awareness.&#x201D; Items 1&#x2013;8 were also deemed usable when assigned to new dimensions, further strengthening the bifactor structure of this scale. Based on these findings, the structural analysis was refined using a model with six dimensions for the EuLeApp&#x00A9; scales and two item-specific dimensions for numbers and images. Six items were removed due to poor model fit, as they could not be assigned to any specific dimension.</p>
<p>As shown in <xref ref-type="fig" rid="fig3">Figure 3</xref>, the IRT curve of the 4-Parameter Logistic (4PL) model is compressed in the y-direction, ensuring that it remains within the probability range defined by the lower bound <inline-formula>
<mml:math id="M2">
<mml:msub>
<mml:mi>&#x03C7;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> and the upper bound <inline-formula>
<mml:math id="M3">
<mml:msub>
<mml:mi>&#x03D2;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> (<xref ref-type="bibr" rid="ref87">Yen et al., 2012</xref>). This structure sets minimum (<inline-formula>
<mml:math id="M4">
<mml:msub>
<mml:mi>&#x03C7;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula>) and maximum <inline-formula>
<mml:math id="M5">
<mml:mfenced open="(" close=")">
<mml:msub>
<mml:mi>&#x03D2;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:mfenced>
</mml:math>
</inline-formula> probabilities for the correct response to each item <inline-formula>
<mml:math id="M6">
<mml:mi>j</mml:mi>
</mml:math>
</inline-formula>, independent of the test-taker&#x2019;s ability level. This allows the model to account for guessing (lower asymptote) and disengagement (upper asymptote), providing a more accurate representation of performance.</p>
<fig position="float" id="fig3">
<label>Figure 3</label>
<caption>
<p>A typical item characteristic curve for the 4PL IRT model. P(<italic>&#x03B8;</italic>) represents the probability of a correct response given the ability level. <inline-formula>
<mml:math id="M7">
<mml:msub>
<mml:mi>&#x03C7;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> (ability). This represents the ability level of the individual, typically ranging from very low to very high values. <inline-formula>
<mml:math id="M8">
<mml:msub>
<mml:mi>&#x03D2;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> (probability of correct response). This represents the probability of answering the item correctly, ranging from 0 to 1. 4PL&#x202F;=&#x202F;four-parameter logistic.</p>
</caption>
<graphic xlink:href="fpsyg-16-1522740-g003.tif"/>
</fig>
<p>While the lower asymptote captures the probability of randomly guessing an item correctly, it is also crucial to consider factors that might cause young children to respond incorrectly to easy items despite possessing the required knowledge. Issues such as inattention, attention deficits, socially desirable responding, external distractions, or a lack of motivation to engage with simple tasks may contribute to this variability (<xref ref-type="bibr" rid="ref004">Liao et al., 2012</xref>). To better accommodate these influences, the 4PL model incorporates an upper asymptote, which accounts for the possibility that even highly skilled individuals may not always respond correctly due to carelessness, disengagement, or momentary lapses in concentration (<xref ref-type="bibr" rid="ref3">Antoniou et al., 2022</xref>; <xref ref-type="bibr" rid="ref005">Barton and Lord, 1981</xref>).</p>
<p>From a developmental perspective, the 4PL model provides a more precise estimate of young children&#x2019;s abilities. By accounting for cognitive and behavioral fluctuations common in early development, this model offers improved sensitivity to response patterns that may be influenced by variability in attention, motivation, and task engagement (<xref ref-type="bibr" rid="ref006">Anderson et al., 2002</xref>). These features make the 4PL model particularly useful in assessing young learners, where performance is not solely determined by ability but also by contextual and developmental factors.</p>
<p>Specific item selection criteria and stopping rules were defined for the item bank used in the Computerized Adaptive Test (CAT) analysis process to enhance testing efficiency and precision (<xref ref-type="bibr" rid="ref26">Ebenbeck and Gebhardt, 2022</xref>). Standard techniques, including the selection of starting items, regression analysis, and stopping rules, were implemented to optimize these goals (<xref ref-type="bibr" rid="ref68">Roberts et al., 2000</xref>). In EuLeApp&#x00A9;, item difficulty is adjusted based on children&#x2019;s responses. To identify age-appropriate starting items, 24 of the 183 available items were selected based on content considerations, ensuring that the MCAT process began with items that were neither too easy nor too difficult for each age group. A regression analysis was then conducted to analyze the influence of these starting items on the standard error of measurement for determining personal abilities. The stopping rule was set based on test precision, where the CAT algorithm assesses whether the confidence interval falls within specified limits. When this criterion is met, the algorithm concludes the assessment for that construct. SEM (standard error of measurement) was chosen for the stopping rule because it offers a flexible framework for modeling relationships between observed data (test items) and latent traits (ability levels), evaluating measurement precision, and enhancing testing efficiency (<xref ref-type="bibr" rid="ref74">Sideridis et al., 2018</xref>). During the assessment process, the stopping rule can be applied to end the test when SEM falls below a specified threshold, indicating sufficient precision:</p>
<p>The formula &#x1D446;&#x1D438;&#x1D440; = <inline-formula>
<mml:math id="M9">
<mml:msqrt>
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>R</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>l</mml:mi>
</mml:mrow>
</mml:msqrt>
<mml:mo>=</mml:mo>
</mml:math>
</inline-formula>&#x21D4; &#x1D445;&#x1D452;&#x1D459; = 1 &#x2212; &#x1D446;&#x1D438;&#x1D440;<sup>2</sup> allows for calculating standard measurement errors corresponding to reliabilities of 0.75, 0.80, 0.85, and 0.90, resulting in SEM values of 0.500, 0.447, 0.387, and 0.316, respectively. When the stopping rule SEM&#x202F;&#x003C;&#x202F;0.316 is applied, the procedure maintains a minimum reliability of 0.90 for the estimates of personal abilities, with the flexibility to specify different minimum reliabilities for each dimension (<xref ref-type="bibr" rid="ref6">Bjorner et al., 2007</xref>; <xref ref-type="bibr" rid="ref20">Chalmers, 2016</xref>). Using this approach, SEM ensures that the adaptive assessment is accurate and efficient, balancing the number of test items with the need for reliable measurement of a child&#x2019;s ability. Simulation studies were conducted to evaluate the performance of the CAT algorithm, generating response patterns from simulated test subjects with fixed parameters (<xref ref-type="bibr" rid="ref48">Magis and Ra&#x00EE;che, 2012</xref>). The ability of IRT models to fit depends on the match between the items and the sample (skewed items require larger sample sizes, such as 500&#x2013;1,000), with larger sample sizes providing better results. Given our smaller sample, we repeated parameter estimates multiple times to enhance their stability. We also conducted simulation studies to evaluate item functionality, establishing factor models based on the intended content during the data generation. The EuLeApp&#x00A9; was built on measuring information on the interrelationships among various early literacy dimensions. MIRT models can estimate skills with the categorical factor structure of early literacy components (<xref ref-type="bibr" rid="ref1">Ackerman et al., 2003</xref>). Thus, the analysis process was carried out with multidimensional CAT, which is based on Multidimensional IRT (MIRT) models and allows the simultaneous measurement of more than one dimension (<xref ref-type="bibr" rid="ref71">Segall, 2009</xref>).</p>
</sec>
</sec>
<sec sec-type="results" id="sec15">
<label>3</label>
<title>Results</title>
<sec id="sec16">
<label>3.1</label>
<title>Multidimensional IRT model comparisons</title>
<p>A confirmatory IRT model was developed by assigning the items to six latent dimensions based on intended measurement purposes: concepts of print (Items 1&#x2013;40), print awareness (Items 41&#x2013;59), word awareness (Items 60&#x2013;71), phonological awareness (Items 71&#x2013;100), alphabet knowledge (Items 101&#x2013;146), and early reading (Items 147&#x2013;183).</p>
<p>Next, a covariance matrix was defined for a model with correlated dimensions. The number of parameters was 408 for the M2PL model, 585 for the M3PL model, and 762 for the M4PL model (<xref ref-type="table" rid="tab3">Table 3</xref>). While statistical model fit is important, it should not be the only criterion for model selection; theoretical assumptions about the underlying model should also be considered (<xref ref-type="bibr" rid="ref69">Robitzsch, 2022</xref>). Since we can theoretically substantiate that the IRT model has four parameters, we compared M2PL, M3PL and M4PL models. However, given the relatively small sample size (n&#x202F;=&#x202F;307), it was anticipated that the model estimates of the M4PL models might lack stability (<xref ref-type="bibr" rid="ref85">Wolf et al., 2013</xref>). When parameter estimates are unstable, this suggests the possibility of alternative models with improved parameter estimations (<xref ref-type="bibr" rid="ref69">Robitzsch, 2022</xref>). Therefore, the M2PL, M3PL, and M4PL models were estimated multiple times using the same item allocations to further address stability concerns (e.g., the M4PL model was estimated 17 times).</p>
<table-wrap position="float" id="tab3">
<label>Table 3</label>
<caption>
<p>Model comparison in search of the optimal structure of EuLeApp&#x00A9;.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th/>
<th/>
<th/>
<th/>
<th/>
<th/>
<th/>
<th align="center" valign="top" colspan="2">Confidence interval</th>
<th/>
<th/>
<th/>
<th align="center" valign="top" colspan="4">Items in MCAT<break/>SEM&#x202F;&#x003C;&#x202F;0.447</th>
</tr>
<tr>
<th align="left" valign="top">Model</th>
<th align="center" valign="top">Nr.</th>
<th align="center" valign="top">Number of parameters</th>
<th align="center" valign="top">
<italic>M<sub>2</sub></italic>
</th>
<th align="center" valign="top">
<italic>dF</italic>
</th>
<th align="center" valign="top">
<italic>p</italic>
</th>
<th align="center" valign="top">RMSEA</th>
<th align="center" valign="top">5% LB</th>
<th align="center" valign="top">95% UB</th>
<th align="center" valign="top">SRMSR</th>
<th align="center" valign="top">TLI</th>
<th align="center" valign="top">CFI</th>
<th align="center" valign="top">
<italic>M</italic>
</th>
<th align="center" valign="top">
<italic>Mdn</italic>
</th>
<th align="center" valign="top">
<italic>SD</italic>
</th>
<th align="center" valign="top">%</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">M2PL</td>
<td align="center" valign="top">1</td>
<td align="center" valign="top">408</td>
<td align="center" valign="top">16495.4</td>
<td align="center" valign="top">15,345</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0157</td>
<td align="center" valign="top">0.0134</td>
<td align="center" valign="top">0.0176</td>
<td align="center" valign="top">0.0641</td>
<td align="center" valign="top">0.9929</td>
<td align="center" valign="top">0.9930</td>
<td align="center" valign="bottom">93.4</td>
<td align="center" valign="bottom">77</td>
<td align="center" valign="bottom">44.9</td>
<td align="center" valign="bottom">19</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">2</td>
<td align="center" valign="top">408</td>
<td align="center" valign="top">16520.7</td>
<td align="center" valign="top">15,345</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0158</td>
<td align="center" valign="top">0.0136</td>
<td align="center" valign="top">0.0178</td>
<td align="center" valign="top">0.0641</td>
<td align="center" valign="top">0.9928</td>
<td align="center" valign="top">0.9929</td>
<td align="center" valign="bottom">92.8</td>
<td align="center" valign="bottom">76</td>
<td align="center" valign="bottom">45.2</td>
<td align="center" valign="bottom">19</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">3</td>
<td align="center" valign="top">408</td>
<td align="center" valign="top">16529.9</td>
<td align="center" valign="top">15,345</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0159</td>
<td align="center" valign="top">0.0137</td>
<td align="center" valign="top">0.0178</td>
<td align="center" valign="top">0.0641</td>
<td align="center" valign="top">0.9927</td>
<td align="center" valign="top">0.9928</td>
<td align="center" valign="bottom">92.2</td>
<td align="center" valign="bottom">76</td>
<td align="center" valign="bottom">44.4</td>
<td align="center" valign="bottom">18</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">4</td>
<td align="center" valign="top">408</td>
<td align="center" valign="top">16549.0</td>
<td align="center" valign="top">15,345</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0160</td>
<td align="center" valign="top">0.0138</td>
<td align="center" valign="top">0.0180</td>
<td align="center" valign="top">0.0640</td>
<td align="center" valign="top">0.9926</td>
<td align="center" valign="top">0.9927</td>
<td align="center" valign="bottom">92.8</td>
<td align="center" valign="bottom">76</td>
<td align="center" valign="bottom">45.0</td>
<td align="center" valign="bottom">19</td>
</tr>
<tr>
<td align="left" valign="top">M3PL</td>
<td align="center" valign="top">1</td>
<td align="center" valign="top">585</td>
<td align="center" valign="top">16352.0</td>
<td align="center" valign="top">15,168</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0160</td>
<td align="center" valign="top">0.0137</td>
<td align="center" valign="top">0.0179</td>
<td align="center" valign="top">0.0679</td>
<td align="center" valign="top">0.9926</td>
<td align="center" valign="top">0.9928</td>
<td align="center" valign="bottom">71.5</td>
<td align="center" valign="bottom">36</td>
<td align="center" valign="bottom">62.9</td>
<td align="center" valign="bottom">25</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">2</td>
<td align="center" valign="top">585</td>
<td align="center" valign="top">16454.5</td>
<td align="center" valign="top">15,168</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0166</td>
<td align="center" valign="top">0.0145</td>
<td align="center" valign="top">0.0185</td>
<td align="center" valign="top">0.0667</td>
<td align="center" valign="top">0.9920</td>
<td align="center" valign="top">0.9922</td>
<td align="center" valign="bottom">67.5</td>
<td align="center" valign="bottom">36</td>
<td align="center" valign="bottom">59.7</td>
<td align="center" valign="bottom">21</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">3</td>
<td align="center" valign="top">585</td>
<td align="center" valign="top">16792.0</td>
<td align="center" valign="top">15,168</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0187</td>
<td align="center" valign="top">0.0168</td>
<td align="center" valign="top">0.0204</td>
<td align="center" valign="top">0.0614</td>
<td align="center" valign="top">0.9899</td>
<td align="center" valign="top">0.9902</td>
<td align="center" valign="bottom">86.2</td>
<td align="center" valign="bottom">48</td>
<td align="center" valign="bottom">66.2</td>
<td align="center" valign="bottom">33</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">4</td>
<td align="center" valign="top">585</td>
<td align="center" valign="top">16946.4</td>
<td align="center" valign="top">15,168</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0196</td>
<td align="center" valign="top">0.0177</td>
<td align="center" valign="top">0.0212</td>
<td align="center" valign="top">0.0594</td>
<td align="center" valign="top">0.9889</td>
<td align="center" valign="top">0.9892</td>
<td align="center" valign="bottom">72.4</td>
<td align="center" valign="bottom">34</td>
<td align="center" valign="bottom">63.9</td>
<td align="center" valign="bottom">25</td>
</tr>
<tr>
<td align="left" valign="top">M4PL</td>
<td align="center" valign="top">1</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">15902.6</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0141</td>
<td align="center" valign="top">0.0116</td>
<td align="center" valign="top">0.0163</td>
<td align="center" valign="top">0.0559</td>
<td align="center" valign="top">0.9943</td>
<td align="center" valign="top">0.9945</td>
<td align="center" valign="top">62.6</td>
<td align="center" valign="top">28</td>
<td align="center" valign="top">63.9</td>
<td align="center" valign="top">23</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">2</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16000.5</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0148</td>
<td align="center" valign="top">0.0124</td>
<td align="center" valign="top">0.0169</td>
<td align="center" valign="top">0.0609</td>
<td align="center" valign="top">0.9936</td>
<td align="center" valign="top">0.9939</td>
<td align="center" valign="top">55.3</td>
<td align="center" valign="top">26</td>
<td align="center" valign="top">58.5</td>
<td align="center" valign="top">18</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">3</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16086.4</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0155</td>
<td align="center" valign="top">0.0131</td>
<td align="center" valign="top">0.0175</td>
<td align="center" valign="top">0.0690</td>
<td align="center" valign="top">0.9931</td>
<td align="center" valign="top">0.9934</td>
<td align="center" valign="top">48.3</td>
<td align="center" valign="top">25</td>
<td align="center" valign="top">55.7</td>
<td align="center" valign="top">15</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">4</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16106.5</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0156</td>
<td align="center" valign="top">0.0133</td>
<td align="center" valign="top">0.0176</td>
<td align="center" valign="top">0.0677</td>
<td align="center" valign="top">0.9930</td>
<td align="center" valign="top">0.9932</td>
<td align="center" valign="top">83.7</td>
<td align="center" valign="top">34</td>
<td align="center" valign="top">71.1</td>
<td align="center" valign="top">36</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">5</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16169.9</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0160</td>
<td align="center" valign="top">0.0138</td>
<td align="center" valign="top">0.0180</td>
<td align="center" valign="top">0.0578</td>
<td align="center" valign="top">0.9926</td>
<td align="center" valign="top">0.9929</td>
<td align="center" valign="top">91.6</td>
<td align="center" valign="top">51</td>
<td align="center" valign="top">73.7</td>
<td align="center" valign="top">42</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">6</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16215.1</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0163</td>
<td align="center" valign="top">0.0141</td>
<td align="center" valign="top">0.0183</td>
<td align="center" valign="top">0.0600</td>
<td align="center" valign="top">0.9923</td>
<td align="center" valign="top">0.9926</td>
<td align="center" valign="top">50.4</td>
<td align="center" valign="top">26</td>
<td align="center" valign="top">55.8</td>
<td align="center" valign="top">16</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">7</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16249.1</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0166</td>
<td align="center" valign="top">0.0144</td>
<td align="center" valign="top">0.0185</td>
<td align="center" valign="top">0.0594</td>
<td align="center" valign="top">0.9921</td>
<td align="center" valign="top">0.9924</td>
<td align="center" valign="top">78.5</td>
<td align="center" valign="top">40</td>
<td align="center" valign="top">64.7</td>
<td align="center" valign="top">29</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">8</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16255.9</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0166</td>
<td align="center" valign="top">0.0145</td>
<td align="center" valign="top">0.0185</td>
<td align="center" valign="top">0.0615</td>
<td align="center" valign="top">0.9920</td>
<td align="center" valign="top">0.9923</td>
<td align="center" valign="top">50.6</td>
<td align="center" valign="top">23</td>
<td align="center" valign="top">59.6</td>
<td align="center" valign="top">18</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">9</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16394.1</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0175</td>
<td align="center" valign="top">0.0154</td>
<td align="center" valign="top">0.0193</td>
<td align="center" valign="top">0.0595</td>
<td align="center" valign="top">0.9912</td>
<td align="center" valign="top">0.9915</td>
<td align="center" valign="top">67.1</td>
<td align="center" valign="top">29</td>
<td align="center" valign="top">66.4</td>
<td align="center" valign="top">26</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">10</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16515.6</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0182</td>
<td align="center" valign="top">0.0163</td>
<td align="center" valign="top">0.0200</td>
<td align="center" valign="top">0.0563</td>
<td align="center" valign="top">0.9904</td>
<td align="center" valign="top">0.9908</td>
<td align="center" valign="top">36.5</td>
<td align="center" valign="top">21</td>
<td align="center" valign="top">45.4</td>
<td align="center" valign="top">9</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">11</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16541.2</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0184</td>
<td align="center" valign="top">0.0164</td>
<td align="center" valign="top">0.0201</td>
<td align="center" valign="top">0.0759</td>
<td align="center" valign="top">0.9902</td>
<td align="center" valign="top">0.9906</td>
<td align="center" valign="top">63.9</td>
<td align="center" valign="top">30</td>
<td align="center" valign="top">63.8</td>
<td align="center" valign="top">23</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">12</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16544.2</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0184</td>
<td align="center" valign="top">0.0164</td>
<td align="center" valign="top">0.0202</td>
<td align="center" valign="top">0.0690</td>
<td align="center" valign="top">0.9902</td>
<td align="center" valign="top">0.9906</td>
<td align="center" valign="top">47.8</td>
<td align="center" valign="top">20</td>
<td align="center" valign="top">59.8</td>
<td align="center" valign="top">17</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">13</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16594.4</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0187</td>
<td align="center" valign="top">0.0168</td>
<td align="center" valign="top">0.0204</td>
<td align="center" valign="top">0.0694</td>
<td align="center" valign="top">0.9899</td>
<td align="center" valign="top">0.9903</td>
<td align="center" valign="top">47.8</td>
<td align="center" valign="top">24</td>
<td align="center" valign="top">55.8</td>
<td align="center" valign="top">15</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">14</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16659.3</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0191</td>
<td align="center" valign="top">0.0172</td>
<td align="center" valign="top">0.0208</td>
<td align="center" valign="top">0.0626</td>
<td align="center" valign="top">0.9895</td>
<td align="center" valign="top">0.9899</td>
<td align="center" valign="top">53.4</td>
<td align="center" valign="top">28</td>
<td align="center" valign="top">56.0</td>
<td align="center" valign="top">15</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">15</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16672.3</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0191</td>
<td align="center" valign="top">0.0173</td>
<td align="center" valign="top">0.0208</td>
<td align="center" valign="top">0.0786</td>
<td align="center" valign="top">0.9894</td>
<td align="center" valign="top">0.9898</td>
<td align="center" valign="top">32.0</td>
<td align="center" valign="top">20</td>
<td align="center" valign="top">40.9</td>
<td align="center" valign="top">7</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">16</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16676.2</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0192</td>
<td align="center" valign="top">0.0173</td>
<td align="center" valign="top">0.0209</td>
<td align="center" valign="top">0.0592</td>
<td align="center" valign="top">0.9894</td>
<td align="center" valign="top">0.9898</td>
<td align="center" valign="top">72.9</td>
<td align="center" valign="top">36</td>
<td align="center" valign="top">68.2</td>
<td align="center" valign="top">29</td>
</tr>
<tr>
<td/>
<td align="center" valign="top">17</td>
<td align="center" valign="top">762</td>
<td align="center" valign="top">16917.5</td>
<td align="center" valign="top">14,991</td>
<td align="center" valign="top">0.000</td>
<td align="center" valign="top">0.0205</td>
<td align="center" valign="top">0.0187</td>
<td align="center" valign="top">0.0221</td>
<td align="center" valign="top">0.0796</td>
<td align="center" valign="top">0.9879</td>
<td align="center" valign="top">0.9883</td>
<td align="center" valign="top">32.0</td>
<td align="center" valign="top">20</td>
<td align="center" valign="top">40.9</td>
<td align="center" valign="top">7</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>
<italic>M&#x2082;, Deviance; dF, Degrees of Freedom; RMSEA, Root Mean Square Error of Approximation; 5% LB, %5 Lower Bound; 95% UB, 95% Upper Bound; SRMSR, Standardized Root Mean Square Residual; TLI &#x0026; CFI, Tucker-Lewis Index &#x0026; Comparative Fit Index; M, Mean Number of Items in MCAT; Mdn, Median Number of Items in MCAT; SD, Standard Deviation; %, Percentage of individuals required to answer all 177 items.</italic>
</p>
</table-wrap-foot>
</table-wrap>
<p><xref ref-type="table" rid="tab3">Table 3</xref> presents the model fit indices for the M2PL, M3PL, and M4PL models, with all estimates using the final item assignments.</p>
<p>The results consistently indicated that M4PL models provided a superior data fit compared to M2PL and M3PL models. Specifically, M4PL models No. 1 and 2 demonstrated superior RMSEA values, while models No. 1 and 10 stood out in terms of their SRMSR values. Based on the literature by <xref ref-type="bibr" rid="ref50">Maydeu-Olivares (2013)</xref> to prioritize SRMSR as a model fit criterion, M4PL model No. 10 was ultimately selected. Additionally, the performance of model No. 10 in the multidimensional CAT framework played a decisive role in its selection. The favored M4PL model is superior to the M2PL and M3PL models not only in terms of fit indices, but also in terms of the number of answered items required to achieve a reliability of 0.80 on the MCAT. This iterative estimation process was essential for ensuring parameter stability and validating the reliability of the adaptive testing framework, reinforcing the robustness of the final model.</p>
<p><xref ref-type="table" rid="tab3">Table 3</xref> shows that M4PL-Model No. 10, based on the M4PL model, demonstrates a very strong model fit, as indicated by the following fit indices: RMSEA&#x202F;=&#x202F;0.0182 [0.0163; 0.0200], which suggests a close fit to the data; SRMR&#x202F;=&#x202F;0.0563, indicating a small standardized residual; TLI&#x202F;=&#x202F;0.9904, and CFI&#x202F;=&#x202F;0.9908, both of which suggest an excellent fit to the model.</p>
<p>These values collectively confirm that the model fits the observed data well, providing reliable estimates.</p>
</sec>
<sec id="sec17">
<label>3.2</label>
<title>MCAT calibration of the item pool</title>
<p>An optimally calibrated item pool for CAT should include a broad distribution of items covering the difficulty parameter range <inline-formula>
<mml:math id="M10">
<mml:msub>
<mml:mi mathvariant="normal">&#x0394;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula>, from &#x2212;2 to +2. This range ensures that the item pool can assess abilities across a wide spectrum of test-takers. In addition to a well-distributed difficulty range, the discrimination parameter <inline-formula>
<mml:math id="M11">
<mml:msub>
<mml:mi>A</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula>, plays a critical role in item selection. Higher values of <inline-formula>
<mml:math id="M12">
<mml:msub>
<mml:mi>A</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> indicate items that are more effective in distinguishing between individuals of different ability levels, thus contributing to the overall reliability of the test. <xref ref-type="fig" rid="fig4">Figure 4</xref> presents the distribution of the difficulty parameters <inline-formula>
<mml:math id="M13">
<mml:msub>
<mml:mi mathvariant="normal">&#x0394;</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> (<italic>MDIFF</italic> on the x-axis) and the discrimination parameters <inline-formula>
<mml:math id="M14">
<mml:msub>
<mml:mi>A</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> (<italic>MDISC</italic> on the y-axis) for the 177 items in the item bank.</p>
<fig position="float" id="fig4">
<label>Figure 4</label>
<caption>
<p>Item difficulties in the calibrated item bank. <italic>MDISC</italic>: Item Discrimination Index, <italic>MDIFF</italic>: Item Difficulty Index. The <italic>MDIFF</italic> values were multiplied by &#x2212;1 so that the difficulty of the items increases from left to right in the diagram.</p>
</caption>
<graphic xlink:href="fpsyg-16-1522740-g004.tif"/>
</fig>
<p>Regarding discrimination <inline-formula>
<mml:math id="M15">
<mml:msub>
<mml:mi>A</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula>, items in the calibrated bank exhibit varying levels of effectiveness in distinguishing between children with different abilities. The majority of the items falls within the range of high discriminatory power (<inline-formula>
<mml:math id="M16">
<mml:msub>
<mml:mi>A</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula>&#x003E; 1.5, dashed line), indicating that they are very effective in distinguishing between different levels of ability. Only a small subset of 16 items shows an acceptable discriminatory power of 0.5 &#x003C; <inline-formula>
<mml:math id="M17">
<mml:msub>
<mml:mi>A</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> &#x003C; 1.5 (area between the dotted and dashed line). In the area of average abilities (&#x2212;1&#x202F;&#x003C;&#x202F;&#x0394;j&#x202F;&#x003C;&#x202F;1), there is a subset of 15 items that stand out with an excellent selectivity of <inline-formula>
<mml:math id="M18">
<mml:msub>
<mml:mi>A</mml:mi>
<mml:mi>j</mml:mi>
</mml:msub>
</mml:math>
</inline-formula>&#x003E;15. These items are especially valuable in increasing the precision of the ability estimates for individuals whose abilities lie near the population mean. The calibrated item pool reveals that items with the highest discriminatory power are concentrated in the range of average abilities, specifically between &#x2212;1 and +1 on the ability scale. The item bank is concentrated with highly discriminating items from below-average to average abilities (&#x2212;2 to +1), ensuring accurate assessment in this range. A dozen items adequately cover the above-average range (+1 to +3). This distribution ensures the test maintains precision and reliability across a broad spectrum of abilities.</p>
<sec id="sec18">
<label>3.2.1</label>
<title>Results from the simulation study</title>
<p>Given the developmental range of the target sample (ages 4 to 7), it was important to examine whether differential item functioning occurs as a result of age. To determine optimal starting items, easy items were selected based on average items and relevant literature per scale for each of the three age groups (4&#x202F;years and younger, 5&#x202F;years, and 6&#x202F;years and older). Approximate Posterior (AP) rule for item selection and the Maximum <italic>A Posteriori</italic> (MAP) method for estimating a person&#x2019;s ability. Each age group was assigned a single starting item, resulting in a total of three starting items. The simulation was conducted twice, applying different stopping rules: SEM&#x202F;&#x003C;&#x202F;0.447 for the first and SEM&#x202F;&#x003C;&#x202F;0.387 for the second. The MCAT simulation with 307 complete cases from the training sample provided standard errors of measurement (SEM) for the estimates of the six early literacy components. The sum of squared SEM values was used as the target variable to be minimized. Then, a linear regression analysis was conducted to assess the influence of the two variables: starting item and stop rule. <xref ref-type="table" rid="tab4">Table 4</xref> shows that age-appropriate starting items significantly improved the ability to estimate precision, with significant reductions in the sum of squared errors per age group: for four-year-olds by 6.92 (&#x2212;1.07 per dimension) (<italic>t</italic>&#x202F;=&#x202F;&#x2212;3.52, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001), for 5-year-olds by 11.05 (&#x2212;1.36 per dimension) (<italic>t</italic>&#x202F;=&#x202F;&#x2212;5.00, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001), and for 6-year-olds by 10.83 (&#x2212;1.34 per dimension) (<italic>t</italic>&#x202F;=&#x202F;&#x2212;4.03, <italic>p</italic>&#x202F;&#x003C;&#x202F;0.001). These findings align with psychometric principles, which emphasize the importance of accounting in assessments covering a broad age range (<xref ref-type="bibr" rid="ref5">Best and Miller, 2010</xref>; <xref ref-type="bibr" rid="ref75">Snow, 2020</xref>).</p>
<table-wrap position="float" id="tab4">
<label>Table 4</label>
<caption>
<p>Regression analysis for identifying optimal starting items by age group.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Age groups</th>
<th align="left" valign="top">Items</th>
<th align="center" valign="top">
<italic>b</italic>
</th>
<th align="center" valign="top">
<italic>SE</italic>
</th>
<th align="center" valign="top">
<italic>t</italic>
</th>
<th align="center" valign="top">
<italic>p</italic>
</th>
<th align="center" valign="top">
<inline-formula>
<mml:math id="M19">
<mml:mi>&#x03B2;</mml:mi>
</mml:math>
</inline-formula>
</th>
</tr>
</thead>
<tbody>
<tr>
<td/>
<td align="left" valign="bottom">(Intercept)</td>
<td align="center" valign="bottom">57.91</td>
<td align="center" valign="bottom">1.74</td>
<td align="center" valign="bottom">33.20</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">0.164</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SEM&#x202F;&#x003C;&#x202F;0.447</td>
<td align="center" valign="bottom">6.19</td>
<td align="center" valign="bottom">0.64</td>
<td align="center" valign="bottom">9.68</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">0.257</td>
</tr>
<tr>
<td align="left" valign="bottom">4;0&#x2013;4;11</td>
<td align="left" valign="bottom">SK01_25boot</td>
<td align="center" valign="bottom">&#x2212;6.92</td>
<td align="center" valign="bottom">1.96</td>
<td align="center" valign="bottom">&#x2212;3.52</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.288</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK06_05mo</td>
<td align="center" valign="bottom">&#x2212;3.23</td>
<td align="center" valign="bottom">2.42</td>
<td align="center" valign="bottom">&#x2212;1.33</td>
<td align="center" valign="bottom">0.184</td>
<td align="center" valign="bottom">&#x2212;0.134</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK06_01buchstaben1s</td>
<td align="center" valign="bottom">&#x2212;2.93</td>
<td align="center" valign="bottom">2.42</td>
<td align="center" valign="bottom">&#x2212;1.21</td>
<td align="center" valign="bottom">0.226</td>
<td align="center" valign="bottom">&#x2212;0.122</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK05_27am</td>
<td align="center" valign="bottom">&#x2212;1.85</td>
<td align="center" valign="bottom">2.42</td>
<td align="center" valign="bottom">&#x2212;0.76</td>
<td align="center" valign="bottom">0.445</td>
<td align="center" valign="bottom">&#x2212;0.077</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK04_16e</td>
<td align="center" valign="bottom">&#x2212;0.91</td>
<td align="center" valign="bottom">2.42</td>
<td align="center" valign="bottom">&#x2212;0.38</td>
<td align="center" valign="bottom">0.708</td>
<td align="center" valign="bottom">&#x2212;0.038</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK01_01einbahnstra&#x00DF;e</td>
<td align="center" valign="bottom">&#x2212;0.55</td>
<td align="center" valign="bottom">2.42</td>
<td align="center" valign="bottom">&#x2212;0.23</td>
<td align="center" valign="bottom">0.820</td>
<td align="center" valign="bottom">&#x2212;0.023</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK02_09buchstabe1</td>
<td align="center" valign="bottom">&#x2212;0.48</td>
<td align="center" valign="bottom">2.10</td>
<td align="center" valign="bottom">&#x2212;0.23</td>
<td align="center" valign="bottom">0.819</td>
<td align="center" valign="bottom">&#x2212;0.020</td>
</tr>
<tr>
<td align="left" valign="bottom">5;0&#x2013;5;11</td>
<td align="left" valign="bottom">SK05_28im</td>
<td align="center" valign="bottom">&#x2212;11.05</td>
<td align="center" valign="bottom">2.21</td>
<td align="center" valign="bottom">&#x2212;5.00</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.459</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK05_12r</td>
<td align="center" valign="bottom">&#x2212;9.08</td>
<td align="center" valign="bottom">2.21</td>
<td align="center" valign="bottom">&#x2212;4.11</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.378</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK03_06luecke</td>
<td align="center" valign="bottom">&#x2212;8.50</td>
<td align="center" valign="bottom">2.21</td>
<td align="center" valign="bottom">&#x2212;3.85</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.353</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK01_03friseur</td>
<td align="center" valign="bottom">&#x2212;8.29</td>
<td align="center" valign="bottom">2.21</td>
<td align="center" valign="bottom">&#x2212;3.75</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.345</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK05_04u</td>
<td align="center" valign="bottom">&#x2212;8.11</td>
<td align="center" valign="bottom">2.21</td>
<td align="center" valign="bottom">&#x2212;3.67</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.337</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK06_06am</td>
<td align="center" valign="bottom">&#x2212;8.02</td>
<td align="center" valign="bottom">2.21</td>
<td align="center" valign="bottom">&#x2212;3.63</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.333</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK04_15f</td>
<td align="center" valign="bottom">&#x2212;7.40</td>
<td align="center" valign="bottom">2.21</td>
<td align="center" valign="bottom">&#x2212;3.35</td>
<td align="center" valign="bottom">0.001</td>
<td align="center" valign="bottom">&#x2212;0.308</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK06_01buchstaben1i</td>
<td align="center" valign="bottom">&#x2212;7.24</td>
<td align="center" valign="bottom">2.21</td>
<td align="center" valign="bottom">&#x2212;3.28</td>
<td align="center" valign="bottom">0.001</td>
<td align="center" valign="bottom">&#x2212;0.301</td>
</tr>
<tr>
<td align="left" valign="bottom">6;0&#x2013;6;11</td>
<td align="left" valign="bottom">SK05_29po</td>
<td align="center" valign="bottom">&#x2212;10.83</td>
<td align="center" valign="bottom">2.69</td>
<td align="center" valign="bottom">&#x2212;4.03</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.450</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK05_19y</td>
<td align="center" valign="bottom">&#x2212;10.18</td>
<td align="center" valign="bottom">2.69</td>
<td align="center" valign="bottom">&#x2212;3.79</td>
<td align="center" valign="bottom">0.000</td>
<td align="center" valign="bottom">&#x2212;0.423</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK02_07seite1</td>
<td align="center" valign="bottom">&#x2212;9.35</td>
<td align="center" valign="bottom">2.69</td>
<td align="center" valign="bottom">&#x2212;3.48</td>
<td align="center" valign="bottom">0.001</td>
<td align="center" valign="bottom">&#x2212;0.389</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK04_29kuh2</td>
<td align="center" valign="bottom">&#x2212;9.00</td>
<td align="center" valign="bottom">2.69</td>
<td align="center" valign="bottom">&#x2212;3.35</td>
<td align="center" valign="bottom">0.001</td>
<td align="center" valign="bottom">&#x2212;0.374</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK06_04buchstaben4q</td>
<td align="center" valign="bottom">&#x2212;8.97</td>
<td align="center" valign="bottom">2.69</td>
<td align="center" valign="bottom">&#x2212;3.34</td>
<td align="center" valign="bottom">0.001</td>
<td align="center" valign="bottom">&#x2212;0.373</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK06_08pa</td>
<td align="center" valign="bottom">&#x2212;8.72</td>
<td align="center" valign="bottom">2.69</td>
<td align="center" valign="bottom">&#x2212;3.25</td>
<td align="center" valign="bottom">0.001</td>
<td align="center" valign="bottom">&#x2212;0.363</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK06_02buchstaben2r</td>
<td align="center" valign="bottom">&#x2212;6.79</td>
<td align="center" valign="bottom">2.69</td>
<td align="center" valign="bottom">&#x2212;2.53</td>
<td align="center" valign="bottom">0.011</td>
<td align="center" valign="bottom">&#x2212;0.282</td>
</tr>
<tr>
<td/>
<td align="left" valign="bottom">SK03_11woerteranzahl3</td>
<td align="center" valign="bottom">&#x2212;6.65</td>
<td align="center" valign="bottom">2.69</td>
<td align="center" valign="bottom">&#x2212;2.47</td>
<td align="center" valign="bottom">0.013</td>
<td align="center" valign="bottom">&#x2212;0.276</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>SK01, concepts of prints; SK02, print awareness; SK03, word awareness; SK04, phonological awareness; SK05, alphabet knowledge; SK06, early reading; <italic>b</italic>, Regression Coefficient; <italic>SE</italic>, Standard Error; <italic>&#x03B2;</italic>, Standardized Coefficient.</p>
<p>The reference group for the dummy-coded regression analysis is the start item SK05_05i, with a stop rule of SEM&#x202F;&#x003C;&#x202F;0.387. Since R<sup>2</sup> is not meaningful in a dummy-coded regression analysis, it is not reported in this table.</p>
</table-wrap-foot>
</table-wrap>
<p>Building on this precision improvement, the next phase focused on determining the most effective item selection method through systematic simulation studies comparing parameter estimates using the 4PL model in mirtCAT (<xref ref-type="bibr" rid="ref20">Chalmers, 2016</xref>) across different item selection rules.</p>
<p>IRT-based CATs employ various rules to estimate children&#x2019;s abilities (<xref ref-type="bibr" rid="ref86">Yao, 2013</xref>). During the assessment process, each response continuously refines the ability estimate, allowing the system to adjust dynamically. The algorithm selects the most appropriate items to optimize ability estimation. In addition, the system adapts based on whether a response is correct or incorrect, adjusting the difficulty level of subsequent items accordingly. Typically, a correct response leads to the selection of a more difficult item, while an incorrect response results in an easier one (<xref ref-type="bibr" rid="ref27">Ebenbeck and Gebhardt, 2024</xref>). <xref ref-type="table" rid="tab5">Table 5</xref> summarizes the results from simulation studies comparing different item selection methods (D-rule, T-rule, A-rule, W-rule, E-rule, TP-rule, AP-rule, WP-rule, and EP-rule) using predetermined start items and applying two different stop-rules: SEM&#x202F;&#x003C;&#x202F;0.447 and SEM&#x202F;&#x003C;&#x202F;0.387.</p>
<table-wrap position="float" id="tab5">
<label>Table 5</label>
<caption>
<p>Simulation studies with different item selection methods.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Stop rule</th>
<th align="center" valign="top">Method</th>
<th align="center" valign="top">D-rule</th>
<th align="center" valign="top">T-rule</th>
<th align="center" valign="top">A-rule</th>
<th align="center" valign="top">W-rule</th>
<th align="center" valign="top">E-rule</th>
<th align="center" valign="top">TP-rule</th>
<th align="center" valign="top">AP-rule</th>
<th align="center" valign="top">WP-rule</th>
<th align="center" valign="top">EP-rule</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top" rowspan="11">SEM&#x202F;&#x003C;&#x202F;0.447</td>
<td align="center" valign="top"><italic>M</italic></td>
<td align="center" valign="top">94.7</td>
<td align="center" valign="top">54.1</td>
<td align="center" valign="top">127.5</td>
<td align="center" valign="top">55.1</td>
<td align="center" valign="top">89.3</td>
<td align="center" valign="top">54.1</td>
<td align="center" valign="top">36.5</td>
<td align="center" valign="top">55.1</td>
<td align="center" valign="top">47.4</td>
</tr>
<tr>
<td align="center" valign="top"><italic>SD</italic></td>
<td align="center" valign="top">37.4</td>
<td align="center" valign="top">50.5</td>
<td align="center" valign="top">40.4</td>
<td align="center" valign="top">50.8</td>
<td align="center" valign="top">55.2</td>
<td align="center" valign="top">50.5</td>
<td align="center" valign="top">45.4</td>
<td align="center" valign="top">50.8</td>
<td align="center" valign="top">54.7</td>
</tr>
<tr>
<td align="center" valign="top">Min</td>
<td align="center" valign="top">71</td>
<td align="center" valign="top">10</td>
<td align="center" valign="top">31</td>
<td align="center" valign="top">11</td>
<td align="center" valign="top">20</td>
<td align="center" valign="top">10</td>
<td align="center" valign="top">9</td>
<td align="center" valign="top">11</td>
<td align="center" valign="top">10</td>
</tr>
<tr>
<td align="center" valign="top">Median</td>
<td align="center" valign="top">79</td>
<td align="center" valign="top">31</td>
<td align="center" valign="top">138</td>
<td align="center" valign="top">32</td>
<td align="center" valign="top">70</td>
<td align="center" valign="top">31</td>
<td align="center" valign="top">21</td>
<td align="center" valign="top">32</td>
<td align="center" valign="top">25</td>
</tr>
<tr>
<td align="center" valign="top">75. Per.</td>
<td align="center" valign="top">83</td>
<td align="center" valign="top">66</td>
<td align="center" valign="top">157</td>
<td align="center" valign="top">62</td>
<td align="center" valign="top">139</td>
<td align="center" valign="top">66</td>
<td align="center" valign="top">29</td>
<td align="center" valign="top">62</td>
<td align="center" valign="top">35</td>
</tr>
<tr>
<td align="center" valign="top">80. Per.</td>
<td align="center" valign="top">86</td>
<td align="center" valign="top">81</td>
<td align="center" valign="top">161</td>
<td align="center" valign="top">74</td>
<td align="center" valign="top">144</td>
<td align="center" valign="top">81</td>
<td align="center" valign="top">31</td>
<td align="center" valign="top">74</td>
<td align="center" valign="top">42</td>
</tr>
<tr>
<td align="center" valign="top">85. Per.</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">110</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">111</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">110</td>
<td align="center" valign="top">37</td>
<td align="center" valign="top">111</td>
<td align="center" valign="top">81</td>
</tr>
<tr>
<td align="center" valign="top">90. Per.</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">54</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">95. Per.</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">Max</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">Total 177 items</td>
<td align="center" valign="top">17%</td>
<td align="center" valign="top">11%</td>
<td align="center" valign="top">18%</td>
<td align="center" valign="top">12%</td>
<td align="center" valign="top">16%</td>
<td align="center" valign="top">11%</td>
<td align="center" valign="top">9%</td>
<td align="center" valign="top">12%</td>
<td align="center" valign="top">15%</td>
</tr>
<tr>
<td align="left" valign="top" rowspan="11">SEM&#x202F;&#x003C;&#x202F;0.387</td>
<td align="center" valign="top"><italic>M</italic></td>
<td align="center" valign="top">112.1</td>
<td align="center" valign="top">86.6</td>
<td align="center" valign="top">144.9</td>
<td align="center" valign="top">87.9</td>
<td align="center" valign="top">117.7</td>
<td align="center" valign="top">86.6</td>
<td align="center" valign="top">60.4</td>
<td align="center" valign="top">87.9</td>
<td align="center" valign="top">77.3</td>
</tr>
<tr>
<td align="center" valign="top"><italic>SD</italic></td>
<td align="center" valign="top">46.7</td>
<td align="center" valign="top">64.8</td>
<td align="center" valign="top">36.4</td>
<td align="center" valign="top">64.2</td>
<td align="center" valign="top">54.4</td>
<td align="center" valign="top">64.8</td>
<td align="center" valign="top">62.3</td>
<td align="center" valign="top">64.2</td>
<td align="center" valign="top">67.4</td>
</tr>
<tr>
<td align="center" valign="top">Min.</td>
<td align="center" valign="top">71</td>
<td align="center" valign="top">14</td>
<td align="center" valign="top">31</td>
<td align="center" valign="top">15</td>
<td align="center" valign="top">23</td>
<td align="center" valign="top">14</td>
<td align="center" valign="top">9</td>
<td align="center" valign="top">15</td>
<td align="center" valign="top">12</td>
</tr>
<tr>
<td align="center" valign="top">Median</td>
<td align="center" valign="top">81</td>
<td align="center" valign="top">56</td>
<td align="center" valign="top">154</td>
<td align="center" valign="top">56</td>
<td align="center" valign="top">132</td>
<td align="center" valign="top">56</td>
<td align="center" valign="top">29</td>
<td align="center" valign="top">56</td>
<td align="center" valign="top">40</td>
</tr>
<tr>
<td align="center" valign="top">75. Per</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">57</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">80. Per</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">85. Per</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">90. Per</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">95. Per</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">Max</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
<td align="center" valign="top">177</td>
</tr>
<tr>
<td align="center" valign="top">Total 177 items</td>
<td align="center" valign="top">34%</td>
<td align="center" valign="top">30%</td>
<td align="center" valign="top">39%</td>
<td align="center" valign="top">30%</td>
<td align="center" valign="top">34%</td>
<td align="center" valign="top">30%</td>
<td align="center" valign="top">21%</td>
<td align="center" valign="top">30%</td>
<td align="center" valign="top">31%</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>SEM, Standard Error of Measurement; M, Mean; SD, Standard Deviation; D-rule, Discrimination Rule; T-rule, Theta Rule; A-rule, A-optimality Rule; W-rule, Weighted Information Rule; E-rule, Entropy Rule; TP-rule, Targeting Precision Rule; WP-rule, Weighted Posterior Rule; EP-rule, Expected Posterior Rule.</p>
</table-wrap-foot>
</table-wrap>
<p>When applying the stop-rule SEM&#x202F;&#x003C;&#x202F;0.447, which corresponds to a minimum reliability of 0.80, the analysis revealed that a minimum of 9 items was required, with a mean of 36.5 items and a median of 21 items for the calibration sample; notably, 90% of the tests were completed within 54 items. With a stricter stopping criterion SEM&#x202F;&#x003C;&#x202F;0.387 (reliability&#x202F;=&#x202F;0.85), the AP rule again performed best, resulting in 75% of tests being completed within 57 items. The AP-rule showed the most optimal item selection method, requiring fewer items to achieve the desired level of precision, as demonstrated by both stopping rules (SEM&#x202F;&#x003C;&#x202F;0.447 and SEM&#x202F;&#x003C;&#x202F;0.387). Based on these findings, the following stop-rule strategy was developed: (a) the first 49 items are evaluated using SEM&#x202F;&#x003C;&#x202F;0.387, ensuring that approximately 75% of cases are tested with a minimum reliability of 0.85. From the 50th item, the stop rule switches to SEM&#x202F;&#x003C;&#x202F;0.447, covering an additional 15% of cases with a minimum reliability of 0.80. Beyond the 60th item, the remaining 10% of cases&#x2014;those requiring all 177 items&#x2014;are addressed. In these cases, the procedure calculates the sum of squares of the standard error ranges for the six dimensions over the last 10 items, terminating when the sum falls below 0.0005. An overview of the stop rules applied across the six Early Literacy dimensions is presented below:<list list-type="bullet">
<list-item>
<p>Stop rule for items 1&#x2013;49: SEM&#x202F;&#x003C;&#x202F;0.387298334620742</p>
</list-item>
<list-item>
<p>Stop rule from item 50 onwards: SEM&#x202F;&#x003C;&#x202F;0.447213595499958</p>
</list-item>
<list-item>
<p>Additional stop rule from item 60: <inline-formula>
<mml:math id="M20">
<mml:munderover>
<mml:mstyle displaystyle="true">
<mml:mo stretchy="true">&#x2211;</mml:mo>
</mml:mstyle>
<mml:mrow>
<mml:mi>f</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mn>6</mml:mn>
</mml:munderover>
<mml:msup>
<mml:mfenced open="(" close=")">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>E</mml:mi>
<mml:msubsup>
<mml:mi>M</mml:mi>
<mml:mi>f</mml:mi>
<mml:mtext>max</mml:mtext>
</mml:msubsup>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mi>E</mml:mi>
<mml:msubsup>
<mml:mi>M</mml:mi>
<mml:mi>f</mml:mi>
<mml:mtext>min</mml:mtext>
</mml:msubsup>
</mml:mrow>
</mml:mfenced>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x003C;</mml:mo>
<mml:mn>0.0005</mml:mn>
</mml:math>
</inline-formula></p>
</list-item>
</list></p>
<p>If the threshold of 0.0005 is met, it indicates that the standard errors of measurement over the last 10 items differ by less than <inline-formula>
<mml:math id="M21">
<mml:msqrt>
<mml:mrow>
<mml:mn>0.0005</mml:mn>
<mml:mo stretchy="true">/</mml:mo>
<mml:mn>6</mml:mn>
</mml:mrow>
</mml:msqrt>
<mml:mo>=</mml:mo>
</mml:math>
</inline-formula>0.01 per dimension. This shows that the standard errors are stable, and cannot be significantly reduced by additional items. According to the stopping rules developed using the calibrated data (<xref ref-type="table" rid="tab6">Table 6</xref>), 75% of children are tested with up to 50 items at a minimum reliability of 0.85, while the remaining cases are tested with up to approximately 75 items at a minimum reliability of 0.80.</p>
<table-wrap position="float" id="tab6">
<label>Table 6</label>
<caption>
<p>Results of MCAT simulations for the number of items per test and the reliability obtained for each dimension.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Sample</th>
<th/>
<th align="center" valign="top">Items</th>
<th align="center" valign="top">Rel. CP</th>
<th align="center" valign="top">Rel. PAW</th>
<th align="center" valign="top">Rel. WA</th>
<th align="center" valign="top">Rel. PA</th>
<th align="center" valign="top">Rel. AK</th>
<th align="center" valign="top">Rel. FR</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top" rowspan="9"><italic>N</italic>&#x202F;=&#x202F;307</td>
<td align="center" valign="top">Mean</td>
<td align="center" valign="top">35.0</td>
<td align="center" valign="top">0.897</td>
<td align="center" valign="top">0.929</td>
<td align="center" valign="top">0.944</td>
<td align="center" valign="top">0.920</td>
<td align="center" valign="top">0.939</td>
<td align="center" valign="top">0.869</td>
</tr>
<tr>
<td align="center" valign="top"><italic>SD</italic></td>
<td align="center" valign="top">18.8</td>
<td align="center" valign="top">0.064</td>
<td align="center" valign="top">0.051</td>
<td align="center" valign="top">0.037</td>
<td align="center" valign="top">0.045</td>
<td align="center" valign="top">0.059</td>
<td align="center" valign="top">0.065</td>
</tr>
<tr>
<td align="center" valign="top">Min</td>
<td align="center" valign="top">9</td>
<td align="center" valign="top">0.558</td>
<td align="center" valign="top">0.619</td>
<td align="center" valign="top">0.682</td>
<td align="center" valign="top">0.659</td>
<td align="center" valign="top">0.578</td>
<td align="center" valign="top">0.522</td>
</tr>
<tr>
<td align="center" valign="top">5. Per.</td>
<td align="center" valign="top">15</td>
<td align="center" valign="top">0.804</td>
<td align="center" valign="top">0.862</td>
<td align="center" valign="top">0.899</td>
<td align="center" valign="top">0.855</td>
<td align="center" valign="top">0.804</td>
<td align="center" valign="top">0.751</td>
</tr>
<tr>
<td align="center" valign="top">25. Per.</td>
<td align="center" valign="top">21</td>
<td align="center" valign="top">0.863</td>
<td align="center" valign="top">0.899</td>
<td align="center" valign="top">0.925</td>
<td align="center" valign="top">0.894</td>
<td align="center" valign="top">0.938</td>
<td align="center" valign="top">0.853</td>
</tr>
<tr>
<td align="center" valign="top">50. Per.</td>
<td align="center" valign="top">29</td>
<td align="center" valign="top">0.913</td>
<td align="center" valign="top">0.928</td>
<td align="center" valign="top">0.945</td>
<td align="center" valign="top">0.924</td>
<td align="center" valign="top">0.960</td>
<td align="center" valign="top">0.864</td>
</tr>
<tr>
<td align="center" valign="top">75. Per.</td>
<td align="center" valign="top">50</td>
<td align="center" valign="top">0.942</td>
<td align="center" valign="top">0.973</td>
<td align="center" valign="top">0.973</td>
<td align="center" valign="top">0.958</td>
<td align="center" valign="top">0.973</td>
<td align="center" valign="top">0.893</td>
</tr>
<tr>
<td align="center" valign="top">95. Per.</td>
<td align="center" valign="top">73</td>
<td align="center" valign="top">0.965</td>
<td align="center" valign="top">0.988</td>
<td align="center" valign="top">0.984</td>
<td align="center" valign="top">0.974</td>
<td align="center" valign="top">0.981</td>
<td align="center" valign="top">0.980</td>
</tr>
<tr>
<td align="center" valign="top">Max</td>
<td align="center" valign="top">118</td>
<td align="center" valign="top">0.979</td>
<td align="center" valign="top">0.989</td>
<td align="center" valign="top">0.989</td>
<td align="center" valign="top">0.982</td>
<td align="center" valign="top">0.993</td>
<td align="center" valign="top">0.989</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>According to the stopping rules developed with the calibrated data, 75% of cases are tested with up to 50 items at a minimum reliability of 0.85, and the remainder are tested with up to approximately 75 items at a minimum reliability of 0.80. For the EuLeApp&#x00A9; dimensions of concepts of print (CP), print awareness (PAW), word awareness (WA), phonological awareness (PA), alphabet knowledge (AK), and first reading (FR), the actual reliabilities are between 0.804 (5th percentile) and 0.988 (95th percentile). It is noticeable that the reliability of the first EuLeApp&#x00A9; dimension of concepts of prints does not reach the reliability of the other five dimensions.</p>
</table-wrap-foot>
</table-wrap>
<p><xref ref-type="fig" rid="fig5">Figure 5</xref> shows the effectiveness of the stop rules and item selection methods, where most participants did not need to complete the maximum number of items. Accordingly, in the distribution of the number of items per early literacy scale, it is evident that tests with up to 50 items form a distinct population.</p>
<fig position="float" id="fig5">
<label>Figure 5</label>
<caption>
<p>Distribution of item selection per test based on stop rule. The figure shows the distribution of the number of items completed (x-axis labeled &#x201C;number of items&#x201D;) and their frequency (y-axis labeled &#x201C;frequency&#x201D;).</p>
</caption>
<graphic xlink:href="fpsyg-16-1522740-g005.tif"/>
</fig>
<p>The two minor peaks observed at around 50 and 60 items correspond to the application of stop rules, such as SEM&#x202F;&#x003C;&#x202F;0.447 or SEM&#x202F;&#x003C;&#x202F;0.387, which marks the end of the test for many participants at these points. Items that show adequate fit to a particular IRT model can be assumed to tap into the construct as specified by the model (<xref ref-type="bibr" rid="ref007">Chan et al., 2015</xref>), while items that show poor fit may measure a different dimension that is not captured by the model specified.</p>
<p>The stop-rule effectively controls standard error, allowing for reliable ability estimates around the 50th item in most cases, as shown in <xref ref-type="fig" rid="fig6">Figure 6</xref>. The dimensions display varying convergence rates in T-scores as more items are answered, with most dimensions stabilizing after approximately 50 to 70 items. For example, dimension SB shows slower T-score stabilization with an R-value of 0.558, indicating moderate reliability. However, dimension EL stabilizes more quickly and shows higher reliability with an R-value of 0.816, as evidenced by the narrower confidence bands and consistent T-scores earlier in the item sequence. Similarly, BK demonstrates strong reliability, with an R-value of 0.785.</p>
<fig position="float" id="fig6">
<label>Figure 6</label>
<caption>
<p>Test of <italic>t</italic>-scores across six scales using stop rule when standard errors (SE)&#x202F;&#x003C;&#x202F;0.0005. The solid lines represent the average <italic>t</italic>-scores, while the shaded areas indicate the standard error ranges.</p>
</caption>
<graphic xlink:href="fpsyg-16-1522740-g006.tif"/>
</fig>
</sec>
</sec>
<sec id="sec19">
<label>3.3</label>
<title>Test&#x2013;retest reliability</title>
<p>The stability of all measures over time was examined through a correlation analysis. T1 took place in Fall 2023, followed by T2 in Spring 2024. As shown in <xref ref-type="table" rid="tab7">Table 7</xref>, which presents both within-time correlations and T1-T2 test&#x2013;retest correlations, each measure demonstrated good reliability across the two time points. The strong correlations (<italic>p</italic>&#x202F;&#x003C;&#x202F;0.01) among literacy components suggest that these skills are highly interrelated, with each influencing or aligning closely with the others.</p>
<table-wrap position="float" id="tab7">
<label>Table 7</label>
<caption>
<p>Combined within-time &#x0026; test&#x2013;retest correlations of early literacy skills.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th/>
<th/>
<th/>
<th/>
<th align="center" valign="top" colspan="6">Pearson correlations</th>
<th/>
</tr>
<tr>
<th align="left" valign="top">Variable</th>
<th align="center" valign="top">
<italic>n</italic>
</th>
<th align="center" valign="top">
<italic>M</italic>
</th>
<th align="center" valign="top">
<italic>SD</italic>
</th>
<th align="center" valign="top">1</th>
<th align="center" valign="top">2</th>
<th align="center" valign="top">3</th>
<th align="center" valign="top">4</th>
<th align="center" valign="top">5</th>
<th align="center" valign="top">6</th>
<th align="center" valign="top">T1-T2</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle" colspan="11">T1</td>
</tr>
<tr>
<td align="left" valign="middle">Concept of prints</td>
<td align="center" valign="middle">307</td>
<td align="center" valign="middle">23.91</td>
<td align="center" valign="middle">10.41</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td/>
<td/>
<td/>
<td/>
<td align="center" valign="top">0.76&#x002A;</td>
</tr>
<tr>
<td align="left" valign="middle">Print awareness</td>
<td align="center" valign="middle">307</td>
<td align="center" valign="middle">12.12</td>
<td align="center" valign="middle">3.86</td>
<td align="center" valign="middle">0.61&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td/>
<td/>
<td/>
<td align="center" valign="top">0.71&#x002A;</td>
</tr>
<tr>
<td align="left" valign="middle">Word awareness</td>
<td align="center" valign="middle">307</td>
<td align="center" valign="middle">5.36</td>
<td align="center" valign="middle">3.20</td>
<td align="center" valign="middle">0.61&#x002A;</td>
<td align="center" valign="middle">0.68&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td/>
<td/>
<td align="center" valign="top">0.73&#x002A;</td>
</tr>
<tr>
<td align="left" valign="middle">Phonologic awareness</td>
<td align="center" valign="middle">307</td>
<td align="center" valign="middle">18.91</td>
<td align="center" valign="middle">4.27</td>
<td align="center" valign="middle">0.51&#x002A;</td>
<td align="center" valign="middle">0.53&#x002A;</td>
<td align="center" valign="middle">0.55&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td/>
<td align="center" valign="top">0.64&#x002A;</td>
</tr>
<tr>
<td align="left" valign="middle">Alphabet knowledge</td>
<td align="center" valign="middle">307</td>
<td align="center" valign="middle">19.57</td>
<td align="center" valign="middle">9.14</td>
<td align="center" valign="middle">0.53&#x002A;</td>
<td align="center" valign="middle">0.45&#x002A;</td>
<td align="center" valign="middle">0.48&#x002A;</td>
<td align="center" valign="middle">0.57&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td align="center" valign="top">0.88&#x002A;</td>
</tr>
<tr>
<td align="left" valign="middle">Early reading</td>
<td align="center" valign="middle">307</td>
<td align="center" valign="middle">7.14</td>
<td align="center" valign="middle">9.05</td>
<td align="center" valign="middle">0.49&#x002A;</td>
<td align="center" valign="middle">0.39&#x002A;</td>
<td align="center" valign="middle">0.40&#x002A;</td>
<td align="center" valign="middle">0.53&#x002A;</td>
<td align="center" valign="middle">0.86&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="top">0.77&#x002A;</td>
</tr>
<tr>
<td align="left" valign="middle" colspan="11">T2</td>
</tr>
<tr>
<td align="left" valign="middle">Concept of prints</td>
<td align="center" valign="middle">159</td>
<td align="center" valign="middle">28.26</td>
<td align="center" valign="middle">9.68</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td/>
<td/>
<td/>
<td/>
<td align="center" valign="top">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">Print awareness</td>
<td align="center" valign="middle">159</td>
<td align="center" valign="middle">13.77</td>
<td align="center" valign="middle">3.56</td>
<td align="center" valign="middle">0.64&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td/>
<td/>
<td/>
<td align="center" valign="top">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">Word awareness</td>
<td align="center" valign="middle">159</td>
<td align="center" valign="middle">6.89</td>
<td align="center" valign="middle">3.33</td>
<td align="center" valign="middle">0.66&#x002A;</td>
<td align="center" valign="middle">0.67&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td/>
<td/>
<td align="center" valign="top">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">Phonologic awareness</td>
<td align="center" valign="middle">159</td>
<td align="center" valign="middle">20.77</td>
<td align="center" valign="middle">4.45</td>
<td align="center" valign="middle">0.56&#x002A;</td>
<td align="center" valign="middle">0.56&#x002A;</td>
<td align="center" valign="middle">0.59&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td/>
<td align="center" valign="top">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">Alphabet knowledge</td>
<td align="center" valign="middle">159</td>
<td align="center" valign="middle">21.27</td>
<td align="center" valign="middle">8.61</td>
<td align="center" valign="middle">0.57&#x002A;</td>
<td align="center" valign="middle">0.49&#x002A;</td>
<td align="center" valign="middle">0.61&#x002A;</td>
<td align="center" valign="middle">0.60&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td/>
<td align="center" valign="top">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">Early reading</td>
<td align="center" valign="middle">159</td>
<td align="center" valign="middle">8.61</td>
<td align="center" valign="middle">8.86</td>
<td align="center" valign="middle">0.50&#x002A;</td>
<td align="center" valign="middle">0.42&#x002A;</td>
<td align="center" valign="middle">0.56&#x002A;</td>
<td align="center" valign="middle">0.53&#x002A;</td>
<td align="center" valign="middle">0.84&#x002A;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="top">&#x2013;</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>&#x002A;<italic>p</italic>&#x202F;&#x003C;&#x202F;0.01.</p>
</table-wrap-foot>
</table-wrap>
</sec>
</sec>
<sec sec-type="discussion" id="sec20">
<label>4</label>
<title>Discussion</title>
<p>This study aimed to develop a conceptual framework for adaptive early literacy assessment tools using MCAT based on Item-Response Theory (IRT). Our findings will be discussed concerning two key points: the psychometric properties of the tablet-based assessment tool and its applications in early literacy education.</p>
<p>The results highlight the powerful psychometric properties of the EuLeApp&#x00A9;, which employs a Multidimensional Computerized Adaptive Testing (MCAT) methodology. The current distribution of item completion demonstrates the effectiveness of this adaptive mechanism (<xref ref-type="bibr" rid="ref45">Keuning and Verhoeven, 2008</xref>), with most children completing the assessment after a variable number of items, reflecting its tailored nature. This approach ensures that the assessment dynamically adjusts to each child&#x2019;s ability level, thereby improving the efficiency and precision of early literacy skill measurement. Simulation studies further indicate that the proposed method optimizes item selection, enhancing the accuracy of ability estimates by leveraging the CAT framework, which adjusts item difficulty based on a child&#x2019;s performance. This adaptive approach yields a more individualized assessment, offering greater precision in measuring early literacy skills compared to traditional fixed-length tests. The method also considers test-taking participation in item selection, which is especially beneficial for children with lower ability levels, particularly in low-category (CAT) settings (<xref ref-type="bibr" rid="ref34">Gorgun and Bulut, 2023</xref>). In sum, these findings have significant implications for measurement practices in adaptive testing (<xref ref-type="bibr" rid="ref81">Weiss, 2004</xref>; <xref ref-type="bibr" rid="ref87">Yen et al., 2012</xref>).</p>
<p>Additionally, while computerized adaptive testing (CAT) typically operates more effectively with simpler models such as the 1PL and 2PL, applying the 4PL model in this study presents areas that warrant further exploration. In this study, the 4PL-IRT model gave better results when using item parameters as an item pool and estimating children&#x2019;s ability levels based on the entire test (183 items). As we mentioned before, one reason for this may be that children&#x2019;s attention may decrease after the 10th minute of the test, which lasts approximately 20&#x202F;min. A parameter that accounts for the probability of carelessness, modeling the chance of a correct answer even when the test child has, the 4PL model could better capture situations where children with sufficient ability to answer questions correctly fail due to a lack of attention or motivation. While the 4-Parameter Logistic (4PL) Model offers flexibility by accounting for guessing and the possibility of a high-performing test child not answering correctly, it may show a weak estimation of item difficulty or discrimination used for adaptive testing with Item Response Theory (IRT). However, to mitigate these challenges, many models have been tested in the calibration phase of the item parameters to ensure that the lower and upper asymptotes are estimated accurately and the items are carefully pre-calibrated. In addition, restrictions were applied to the parameter ranges or regularization techniques of the 4PL model to prevent overfitting and enhance the robustness of the parameter estimates. While these methods helped mitigate some potential issues, further research is necessary to clarify these adjustments&#x2019; impact more clearly. Specifically, future studies should explore how these restrictions influence the accuracy and stability of parameter estimates within the 4PL model, particularly in early literacy assessment contexts. Investigating the long-term effects of using the 4PL model across different populations and educational settings will provide greater insight into its practical value. Continued evaluation and iterative development of EuleApp&#x00A9; will allow for a more robust tool that meets the needs of both educators and young learners better, supporting more precise assessments of early literacy skills.</p>
<p>The EuLeApp&#x00A9; assessment tool has the potential to effectively identify children&#x2019;s strengths and areas of weaknesses in early literacy development. Specifically, its adaptive, data-driven, and multidimensional framework may offer advantages over traditional methods in educational settings. The findings indicate that EuLeApp&#x00A9; can tailor assessments to each child&#x2019;s individual ability levels through sophisticated item selection rules. This dynamic adjustment enables more precise and targeted literacy evaluations than traditional assessments, which typically do not modify item difficulty in response to a child&#x2019;s performance during testing. Moreover, the present study has some important implications for screening children early for possible early reading-writing problems and dyslexia. The literature documents the relationship between early literacy skills and reading achievement well (<xref ref-type="bibr" rid="ref008">Whitehurst and Lonigan, 1998</xref>; <xref ref-type="bibr" rid="ref44">Justice et al., 2009</xref>; <xref ref-type="bibr" rid="ref41">Jim&#x00E9;nez et al., 2024</xref>) and highlights the potential for early monitoring of these skills (<xref ref-type="bibr" rid="ref17">Catts et al., 2001</xref>; <xref ref-type="bibr" rid="ref61">Neumann and Neumann, 2014</xref>). Additionally, other research has pointed out that early comprehensive and accurate assessments provide stronger predictions regarding dyslexia risk and later reading problems (<xref ref-type="bibr" rid="ref17">Catts et al., 2001</xref>; <xref ref-type="bibr" rid="ref18">Catts et al., 2015</xref>). However, studies have often been limited by digital assessment tools that focus only on one or two early literacy domains (e.g., <xref ref-type="bibr" rid="ref32">Golinkoff et al., 2017</xref>; Jonathan <xref ref-type="bibr" rid="ref009">Casta&#x00F1;eda-Fern&#x00E1;ndez et al., 2023</xref>; <xref ref-type="bibr" rid="ref59">Neumann, 2018</xref>). Our findings showed that our more comprehensive screening app can identify children at risk for early literacy, which is crucial for later reading comprehension and school success, with acceptable level of accuracy. The potential of the EuleApp&#x00A9; to accurately measure literacy skills suggests that it could serve to identify children who may benefit from early interventions, thus mitigating the risks associated with dyslexia and other literacy challenges. It is important to note that children enter kindergarten with varying levels of language and cognitive skills, which are critical predictors of future reading success. In the standardization process of this assessment tool, our initial focus was on evaluating children with typical language development to establish baseline performance measures. In the next phase, further development of this screening tool is needed to ensure that children&#x2019;s early literacy accurately distinguishes among diverse children. Overall, using the tool in this way in an early educational setting would be useful for both identifying literacy practices in everyday practice integration, enabling teachers to reflect on their practice, and monitoring progress.</p>
<p>In conclusion, the integration of adaptive, data-driven tools like the EuleApp&#x00A9; in educational settings not only addresses the need for precise assessments but also underscores the critical importance of proactive measures in supporting children&#x2019;s literacy development. The results showed that the individual component scores exhibit adequate psychometric properties, including sufficient precision and validity, supporting the reliability of the assessment tool in measuring early literacy skills effectively.</p>
</sec>
<sec id="sec21">
<label>5</label>
<title>Limitations</title>
<p>This study presents several limitations that should be acknowledged. One limitation of this research is the relatively small sample size. A larger number of participants would facilitate more precise calibration of item parameters and enhance the accuracy of ability estimates across diverse groups (<xref ref-type="bibr" rid="ref6">Bjorner et al., 2007</xref>). Although the current sample size was sufficient to conduct the Computerized Adaptive Testing (CAT) analysis, employing a larger sample would improve the generalizability and robustness of the results. The findings of this study may also lack generalizability due to the limited diversity within the sample. This research was predominantly conducted with a German population and in the German language, which may not fully represent the experiences of children from diverse cultural, linguistic, or socioeconomic backgrounds. The composition of the sample, mainly consisting of participants from middle- and high-socioeconomic-status regions, may present a potential limitation regarding the interpretation of item parameters. In other words, due to their exposure to more literacy-rich environments, these children may have found certain items easier, potentially resulting in a ceiling effect. For instance, assessments of early phonemic awareness (<xref ref-type="bibr" rid="ref12">Burt and Barbara Dodd, 1999</xref>), alphabet knowledge, or early reading skills (<xref ref-type="bibr" rid="ref7">Bowey, 1995</xref>; <xref ref-type="bibr" rid="ref24">Dolean et al., 2019</xref>) may have been less discriminatory for children with early literacy advantages.</p>
<p>Future studies should incorporate more varied samples by increasing the sample size and including participants from diverse cultural and socioeconomic backgrounds to enhance the current findings. Ensuring that all young children have the opportunity to develop proficiency in literacy skills is a key area of focus globally. Consequently, further research involving different populations is essential for exploring the broader applicability of the EuleApp&#x00A9; in various linguistic contexts and ensuring its effectiveness across diverse educational settings. In addition, the study&#x2019;s design may limit the depth of insight into the effectiveness of the EuleApp&#x00A9; assessment tool. For instance, a cross-sectional design captures data at a single point in time, which may not adequately reflect the changes in children&#x2019;s literacy skills over time. Longitudinal studies could provide more comprehensive insights into how children&#x2019;s abilities develop and how effectively the tool tracks these changes.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="sec22">
<title>Data availability statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec sec-type="ethics-statement" id="sec23">
<title>Ethics statement</title>
<p>The studies involving humans were approved by the Research Ethics Committee of the Carl von Ossietzky University of Oldenburg. The studies were conducted in accordance with the local legislation and institutional requirements. Written informed consent to participate in this study was provided by the participants or participants&#x2019; legal guardian/next of kin.</p>
</sec>
<sec sec-type="author-contributions" id="sec24">
<title>Author contributions</title>
<p>MY: Conceptualization, Formal analysis, Supervision, Writing &#x2013; original draft, Writing &#x2013; review &#x0026; editing. CS: Investigation, Writing &#x2013; review &#x0026; editing. MM: Funding acquisition, Project administration, Writing &#x2013; review &#x0026; editing. HL: Writing &#x2013; review &#x0026; editing, Methodology, Formal analysis, Visualization, Software. TJ: Conceptualization, Funding acquisition, Project administration, Supervision, Writing &#x2013; review &#x0026; editing.</p>
</sec>
<sec sec-type="funding-information" id="sec25">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research and/or publication of this article. This work was supported by the German Federal Ministry of Education and Research (BMBF) [grant number 01NV2105A].</p>
</sec>
<ack>
<p>We sincerely appreciate and thank all families, teachers, and children who participated in the study. We also thank our research assistants, who helped us collect data and insight during the research.</p>
</ack>
<sec sec-type="COI-statement" id="sec26">
<title>Conflict of interest</title>
<p>HL is the founder of DHL Data Science Seminars.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="sec27">
<title>Generative AI statement</title>
<p>The authors declare that no Gen AI was used in the creation of this manuscript.</p>
</sec>
<sec sec-type="disclaimer" id="sec28">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="ref1"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ackerman</surname> <given-names>T. A.</given-names></name> <name><surname>Gierl</surname> <given-names>M. J.</given-names></name> <name><surname>Walker</surname> <given-names>C. M.</given-names></name></person-group> (<year>2003</year>). <article-title>Using multidimensional item response theory to evaluate educational and psychological tests</article-title>. <source>Educ. Meas. Issues Pract.</source> <volume>22</volume>, <fpage>37</fpage>&#x2013;<lpage>51</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1745-3992.2003.tb00136.x</pub-id></citation></ref>
<ref id="ref2"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Adams</surname> <given-names>M. J.</given-names></name></person-group> (<year>1994</year>). <source>Beginning to read: Thinking and learning about print</source>.</citation></ref>
<ref id="ref3"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Antoniou</surname> <given-names>F.</given-names></name> <name><surname>Ralli</surname> <given-names>A. M.</given-names></name> <name><surname>Mouzaki</surname> <given-names>A.</given-names></name> <name><surname>Diamanti</surname> <given-names>V.</given-names></name> <name><surname>Papaioannou</surname> <given-names>S.</given-names></name></person-group> (<year>2022</year>). <article-title>Logometro&#x00AE;: the psychometric properties of a norm-referenced digital battery for language assessment of Greek-speaking 4&#x2013;7 years old children</article-title>. <source>Front. Psychol.</source> <volume>13</volume>:<fpage>900600</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2022.900600</pub-id>, PMID: <pub-id pub-id-type="pmid">35959077</pub-id></citation></ref>
<ref id="ref006"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Anderson</surname> <given-names>V.</given-names></name> <name><surname>Levin</surname> <given-names>H. S.</given-names></name> <name><surname>Jacobs</surname> <given-names>R.</given-names></name></person-group> (<year>2002</year>). <article-title>Executive functions after frontal lobe injury: A developmental perspective</article-title>. <source>J. Clin. Exp. Neuropsychol.</source> <volume>24</volume>, <fpage>224</fpage>&#x2013;<lpage>247</lpage>. doi: <pub-id pub-id-type="doi">10.1207/S15327647JCD1201_02</pub-id></citation></ref>
<ref id="ref4"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Baker</surname> <given-names>F.</given-names></name></person-group> (<year>2001</year>). <source>The basics of item response theory: ERIC clearinghouse on assessment and evaluation</source>. <publisher-loc>College Park, MD</publisher-loc>: <publisher-name>University of Maryland, ERIC Clearinghouse on Assessment and Evaluation</publisher-name>.</citation></ref>
<ref id="ref005"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Barton</surname> <given-names>M. A.</given-names></name> <name><surname>Lord</surname> <given-names>F. M.</given-names></name></person-group> (<year>1981</year>). <article-title>An upper asymptote for the three-parameter logistic item-response model</article-title>. <source>ETS Research Report Series</source> <volume>1981</volume>. doi: <pub-id pub-id-type="doi">10.1002/j.2333-8504.1981.tb01255.x</pub-id></citation></ref>
<ref id="ref5"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Best</surname> <given-names>J. R.</given-names></name> <name><surname>Miller</surname> <given-names>P. H.</given-names></name></person-group> (<year>2010</year>). <article-title>A developmental perspective on executive function</article-title>. <source>Child Dev.</source> <volume>81</volume>, <fpage>1641</fpage>&#x2013;<lpage>1660</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1467-8624.2010.01499.x</pub-id>, PMID: <pub-id pub-id-type="pmid">21077853</pub-id></citation></ref>
<ref id="ref6"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bjorner</surname> <given-names>J. B.</given-names></name> <name><surname>Chang</surname> <given-names>C. H.</given-names></name> <name><surname>Thissen</surname> <given-names>D.</given-names></name> <name><surname>Reeve</surname> <given-names>B. B.</given-names></name></person-group> (<year>2007</year>). <article-title>Developing tailored instruments: item banking and computerized adaptive assessment</article-title>. <source>Qual. Life Res.</source> <volume>16</volume>, <fpage>95</fpage>&#x2013;<lpage>108</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s11136-007-9168-6</pub-id>, PMID: <pub-id pub-id-type="pmid">17530450</pub-id></citation></ref>
<ref id="ref7"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bowey</surname> <given-names>J. A.</given-names></name></person-group> (<year>1995</year>). <article-title>Socioeconomic status differences in preschool phonological sensitivity and first-grade reading achievement</article-title>. <source>J. Educ. Psychol.</source> <volume>87</volume>, <fpage>476</fpage>&#x2013;<lpage>487</lpage>. doi: <pub-id pub-id-type="doi">10.1037/0022-0663.87.3.476</pub-id></citation></ref>
<ref id="ref8"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Brassard</surname> <given-names>M. R.</given-names></name> <name><surname>Boehm</surname> <given-names>A. E.</given-names></name></person-group> (<year>2007</year>). &#x201C;<article-title>Assessment of emotional development and behavior problems</article-title>&#x201D; in ed. B. A. Bracken. <source>Preschool assessment: principles and practices</source>. <publisher-loc>New York</publisher-loc>: <publisher-name>Guilford Press</publisher-name>. <fpage>508</fpage>&#x2013;<lpage>576</lpage>.</citation></ref>
<ref id="ref9"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Brown</surname> <given-names>A. T.</given-names></name></person-group> (<year>2015</year>). <source>Confirmatory factor analysis for applied research</source>. <edition>2nd</edition> Edn. <publisher-loc>New York</publisher-loc>: <publisher-name>The Guilford Press</publisher-name>.</citation></ref>
<ref id="ref10"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Buckingham</surname> <given-names>J.</given-names></name> <name><surname>Wheldall</surname> <given-names>K.</given-names></name> <name><surname>Beaman-Wheldall</surname> <given-names>R.</given-names></name></person-group> (<year>2013</year>). <article-title>Why poor children are more likely to become poor readers: the school years</article-title>. <source>Aust. J. Educ.</source> <volume>57</volume>, <fpage>190</fpage>&#x2013;<lpage>213</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0004944113495500</pub-id></citation></ref>
<ref id="ref11"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bulut</surname> <given-names>O.</given-names></name> <name><surname>Cormier</surname> <given-names>D. C.</given-names></name></person-group> (<year>2018</year>). <article-title>Validity evidence for progress monitoring with star reading: slope estimates, administration frequency, and number of data points</article-title>. <source>Front. Educ.</source> <volume>3</volume>:<fpage>68</fpage>. doi: <pub-id pub-id-type="doi">10.3389/feduc.2018.00068</pub-id></citation></ref>
<ref id="ref12"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Burt</surname> <given-names>A. H.</given-names></name> <name><surname>Barbara Dodd</surname> <given-names>L.</given-names></name></person-group> (<year>1999</year>). <article-title>Phonological awareness skills of 4-year-old British children: an assessment and developmental data</article-title>. <source>Int. J. Lang. Commun. Disord.</source> <volume>34</volume>, <fpage>311</fpage>&#x2013;<lpage>335</lpage>. doi: <pub-id pub-id-type="doi">10.1080/136828299247432</pub-id>, PMID: <pub-id pub-id-type="pmid">10884904</pub-id></citation></ref>
<ref id="ref003"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Bruder</surname> <given-names>I</given-names></name></person-group>. (<year>1993</year>). <article-title>Alternative assessment: Putting technology to the test. Electronic Learning</article-title>, <volume>12</volume>, <fpage>22</fpage>&#x2013;<lpage>23</lpage>.</citation></ref>
<ref id="ref13"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Campbell</surname> <given-names>C.</given-names></name> <name><surname>Jane</surname> <given-names>B.</given-names></name></person-group> (<year>2012</year>). <article-title>Motivating children to learn: the role of technology education</article-title>. <source>Int. J. Technol. Des. Educ.</source> <volume>22</volume>, <fpage>1</fpage>&#x2013;<lpage>11</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10798-010-9134-4</pub-id></citation></ref>
<ref id="ref14"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Care</surname> <given-names>E.</given-names></name> <name><surname>Kim</surname> <given-names>H.</given-names></name> <name><surname>Vista</surname> <given-names>A.</given-names></name> <name><surname>Anderson</surname> <given-names>K.</given-names></name></person-group> (<year>2018</year>). <source>Education system alignment for 21st century skills: Focus on assessment</source>. <publisher-loc>Washington, DC</publisher-loc>: <publisher-name>Center for Universal Education at The Brookings Institution</publisher-name>.</citation></ref>
<ref id="ref15"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Carson</surname> <given-names>K.</given-names></name> <name><surname>Boustead</surname> <given-names>T.</given-names></name> <name><surname>Gillon</surname> <given-names>G.</given-names></name></person-group> (<year>2015</year>). <article-title>Content validity to support the use of a computer-based phonological awareness screening and monitoring assessment (com PASMA) in the classroom</article-title>. <source>Int. J. Speech Lang. Pathol.</source> <volume>17</volume>, <fpage>500</fpage>&#x2013;<lpage>510</lpage>. doi: <pub-id pub-id-type="doi">10.3109/17549507.2015.1016107</pub-id>, PMID: <pub-id pub-id-type="pmid">25764226</pub-id></citation></ref>
<ref id="ref009"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Casta&#x00F1;eda-Fern&#x00E1;ndez</surname> <given-names>J.</given-names></name> <name><surname>Neira-Pi&#x00F1;eiro</surname> <given-names>M. R.</given-names></name> <name><surname>L&#x00F3;pez-Bouzas</surname> <given-names>N.</given-names></name> <name><surname>del-Moral-P&#x00E9;rez</surname> <given-names>M. E.</given-names></name></person-group> (<year>2023</year>). <article-title>Empirical validation of the Oral Narrative Competence Evaluation with the TellingApp (ONCE) Scale in early childhood</article-title>. <source>Int. J. Child-Comput. Interact.</source> <volume>36</volume>, <fpage>100580</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.ijcci.2023.100580</pub-id></citation></ref>
<ref id="ref16"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Catts</surname> <given-names>H. W.</given-names></name> <name><surname>Fey</surname> <given-names>M. E.</given-names></name> <name><surname>Tomblin</surname> <given-names>J. B.</given-names></name> <name><surname>Zhang</surname> <given-names>X.</given-names></name></person-group> (<year>2002</year>). <article-title>A longitudinal investigation of reading outcomes in children with language impairments</article-title>. <source>J. Speech Lang. Hear. Res.</source> <volume>45</volume>, <fpage>1142</fpage>&#x2013;<lpage>1157</lpage>. doi: <pub-id pub-id-type="doi">10.1044/1092-4388(2002/093)</pub-id></citation></ref>
<ref id="ref17"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Catts</surname> <given-names>H. W.</given-names></name> <name><surname>Fey</surname> <given-names>M. E.</given-names></name> <name><surname>Zhang</surname> <given-names>X.</given-names></name> <name><surname>Tomlin</surname> <given-names>J. B.</given-names></name></person-group> (<year>2001</year>). <article-title>Estimating the risk of future reading difficulties in kindergarten children: a research-based model and its clinical implementation</article-title>. <source>Lang. Speech. Hear. Serv. Sch.</source> <volume>32</volume>, <fpage>38</fpage>&#x2013;<lpage>50</lpage>. doi: <pub-id pub-id-type="doi">10.1044/0161-1461(2001/004)</pub-id></citation></ref>
<ref id="ref18"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Catts</surname> <given-names>H. W.</given-names></name> <name><surname>Nielsen</surname> <given-names>D. C.</given-names></name> <name><surname>Bridges</surname> <given-names>M. S.</given-names></name> <name><surname>Liu</surname> <given-names>Y. S.</given-names></name> <name><surname>Bontempo</surname> <given-names>D. E.</given-names></name></person-group> (<year>2015</year>). <article-title>Early identification of reading disabilities within an RTI framework</article-title>. <source>J. Learn. Disabil.</source> <volume>48</volume>, <fpage>281</fpage>&#x2013;<lpage>297</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0022219413498115</pub-id>, PMID: <pub-id pub-id-type="pmid">23945079</pub-id></citation></ref>
<ref id="ref007"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chan</surname> <given-names>K. S.</given-names></name> <name><surname>Gross</surname> <given-names>A. L.</given-names></name> <name><surname>Pezzin</surname> <given-names>L. E.</given-names></name> <name><surname>Brandt</surname> <given-names>J.</given-names></name> <name><surname>Kasper</surname> <given-names>J. D.</given-names></name></person-group> (<year>2015</year>). <article-title>Harmonizing measures of cognitive performance across international surveys of aging using item response theory</article-title>. <source>J. Aging Health</source>. <volume>27</volume>, <fpage>1392</fpage>&#x2013;<lpage>1414</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0898264315583054</pub-id></citation></ref>
<ref id="ref19"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Chalmers</surname> <given-names>R. P.</given-names></name></person-group> (<year>2015</year>) <italic>mirtCAT: computerized adaptive testing with multidimensional item response theory</italic>. R package version 0.6, 1.</citation></ref>
<ref id="ref20"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chalmers</surname> <given-names>R. P.</given-names></name></person-group> (<year>2016</year>). <article-title>Generating adaptive and non-adaptive test interfaces for multidimensional item response theory applications</article-title>. <source>J. Stat. Softw.</source> <volume>71</volume>, <fpage>1</fpage>&#x2013;<lpage>39</lpage>. doi: <pub-id pub-id-type="doi">10.18637/jss.v071.i05</pub-id></citation></ref>
<ref id="ref21"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chaney</surname> <given-names>C.</given-names></name></person-group> (<year>1994</year>). <article-title>Language development, metalinguistic awareness, and emergent literacy skills of 3-year-old children in relation to social class</article-title>. <source>Appl. Psycholinguist.</source> <volume>15</volume>, <fpage>371</fpage>&#x2013;<lpage>394</lpage>. doi: <pub-id pub-id-type="doi">10.1017/S0142716400004501</pub-id></citation></ref>
<ref id="ref22"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Chen</surname> <given-names>F. F.</given-names></name> <name><surname>West</surname> <given-names>S. G.</given-names></name> <name><surname>Sousa</surname> <given-names>K. H.</given-names></name></person-group> (<year>2006</year>). <article-title>A comparison of bifactor and second-order models of quality of life</article-title>. <source>Multivar. Behav. Res.</source> <volume>41</volume>, <fpage>189</fpage>&#x2013;<lpage>225</lpage>. doi: <pub-id pub-id-type="doi">10.1207/s15327906mbr4102_5</pub-id>, PMID: <pub-id pub-id-type="pmid">26782910</pub-id></citation></ref>
<ref id="ref23"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Davey</surname> <given-names>A.</given-names></name></person-group> (<year>2005</year>). <article-title>Issues in evaluating model fit with missing data</article-title>. <source>Struct. Equ. Model.</source> <volume>12</volume>, <fpage>578</fpage>&#x2013;<lpage>597</lpage>. doi: <pub-id pub-id-type="doi">10.1207/s15328007sem1204_4</pub-id></citation></ref>
<ref id="ref24"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dolean</surname> <given-names>D.</given-names></name> <name><surname>Melby-Lerv&#x00E5;g</surname> <given-names>M.</given-names></name> <name><surname>Tincas</surname> <given-names>I.</given-names></name> <name><surname>Damsa</surname> <given-names>C.</given-names></name> <name><surname>Lerv&#x00E5;g</surname> <given-names>A.</given-names></name></person-group> (<year>2019</year>). <article-title>Achievement gap: socioeconomic status affects reading development beyond language and cognition in children facing poverty</article-title>. <source>Learn. Instr.</source> <volume>63</volume>:<fpage>101218</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.learninstruc.2019.101218</pub-id></citation></ref>
<ref id="ref25"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Dunn</surname> <given-names>K. J.</given-names></name> <name><surname>McCray</surname> <given-names>G.</given-names></name></person-group> (<year>2020</year>). <article-title>The place of the bifactor model in confirmatory factor analysis investigations into construct dimensionality in language testing</article-title>. <source>Front. Psychol.</source> <volume>11</volume>:<fpage>1357</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2020.01357</pub-id>, PMID: <pub-id pub-id-type="pmid">32765335</pub-id></citation></ref>
<ref id="ref26"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ebenbeck</surname> <given-names>N.</given-names></name> <name><surname>Gebhardt</surname> <given-names>M.</given-names></name></person-group> (<year>2022</year>). <article-title>Simulating computerized adaptive testing in special education based on inclusive progress monitoring data</article-title>. <source>Front. Educ.</source> <volume>7</volume>:<fpage>945733</fpage>. doi: <pub-id pub-id-type="doi">10.3389/feduc.2022.945733</pub-id></citation></ref>
<ref id="ref27"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ebenbeck</surname> <given-names>N.</given-names></name> <name><surname>Gebhardt</surname> <given-names>M.</given-names></name></person-group> (<year>2024</year>). <article-title>Differential performance of computerized adaptive testing in students with and without disabilities&#x2013;a simulation study</article-title>. <source>J. Spec. Educ. Technol.</source> <volume>39</volume>, <fpage>481</fpage>&#x2013;<lpage>490</lpage>. doi: <pub-id pub-id-type="doi">10.1177/01626434241232117</pub-id></citation></ref>
<ref id="ref28"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Elimelech</surname> <given-names>A.</given-names></name> <name><surname>Aram</surname> <given-names>D.</given-names></name></person-group> (<year>2020</year>). <article-title>Using a digital spelling game for promoting alphabetic knowledge of preschoolers: the contribution of auditory and visual supports</article-title>. <source>Read. Res. Q.</source> <volume>55</volume>, <fpage>235</fpage>&#x2013;<lpage>250</lpage>. doi: <pub-id pub-id-type="doi">10.1002/rrq.264</pub-id></citation></ref>
<ref id="ref29"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Engel</surname> <given-names>S.</given-names></name></person-group> (<year>1995</year>). <source>The stories children tell: Making sense of the narratives of childhood</source>. <publisher-loc>New York</publisher-loc>: <publisher-name>Henry Holt and Company</publisher-name>.</citation></ref>
<ref id="ref30"><citation citation-type="other"><person-group person-group-type="author"><collab id="coll1">European Commission</collab></person-group> (<year>2023</year>) &#x2018;Germany: Early childhood education and care&#x2019;. Available online at: <ext-link xlink:href="https://eurydice.eacea.ec.europa.eu/national-education-systems/germany/educational-guidelines" ext-link-type="uri">https://eurydice.eacea.ec.europa.eu/national-education-systems/germany/educational-guidelines</ext-link> (Accessed October 30, 2024).</citation></ref>
<ref id="ref31"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Gee</surname> <given-names>J. P.</given-names></name></person-group> (<year>2012</year>). &#x201C;<article-title>What is literacy?</article-title>&#x201D; in eds. <person-group person-group-type="editor"><name><surname>Simpson</surname> <given-names>J.</given-names></name> <name><surname>Mayr</surname> <given-names>A.</given-names></name> <name><surname>Jones</surname> <given-names>E.</given-names></name></person-group>. <source>Language and linguistics in context</source> (<publisher-loc>London</publisher-loc>: <publisher-name>Routledge</publisher-name>), <fpage>257</fpage>&#x2013;<lpage>264</lpage>.</citation></ref>
<ref id="ref32"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Golinkoff</surname> <given-names>R. M.</given-names></name> <name><surname>De Villiers</surname> <given-names>J. G.</given-names></name> <name><surname>Hirsh-Pasek</surname> <given-names>K.</given-names></name> <name><surname>Iglesias</surname> <given-names>A.</given-names></name> <name><surname>Wilson</surname> <given-names>M. S.</given-names></name> <name><surname>Morini</surname> <given-names>G.</given-names></name> <etal/></person-group>. (<year>2017</year>). <source>User's manual for the quick interactive language screener (QUILS): A measure of vocabulary, syntax, and language acquisition skills in young children</source>. <publisher-loc>Baltimore, MD</publisher-loc>: <publisher-name>Paul H. Brookes Publishing Company</publisher-name>.</citation></ref>
<ref id="ref33"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Gonski</surname> <given-names>D.</given-names></name> <name><surname>Arcus</surname> <given-names>T.</given-names></name> <name><surname>Boston</surname> <given-names>K.</given-names></name> <name><surname>Gould</surname> <given-names>V.</given-names></name> <name><surname>Johnson</surname> <given-names>W.</given-names></name> <name><surname>O&#x2019;Brien</surname> <given-names>L.</given-names></name> <etal/></person-group>. (<year>2018</year>). <source>Through growth to achievement: Report of the review to achieve educational excellence in Australian schools</source>. <publisher-loc>Canberra</publisher-loc>: <publisher-name>Commonwealth of Australia</publisher-name>.</citation></ref>
<ref id="ref34"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Gorgun</surname> <given-names>G.</given-names></name> <name><surname>Bulut</surname> <given-names>O.</given-names></name></person-group> (<year>2023</year>). <article-title>Incorporating test-taking engagement into the item selection algorithm in low-stakes computerized adaptive tests</article-title>. <source>Large-Scale Assess. Educ.</source> <volume>11</volume>:<fpage>27</fpage>. doi: <pub-id pub-id-type="doi">10.1186/s40536-023-00177-5</pub-id></citation></ref>
<ref id="ref35"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Haladyna</surname> <given-names>T. M.</given-names></name></person-group> (<year>2013</year>). <source>Developing and validating test items</source>. <publisher-loc>London</publisher-loc>: <publisher-name>Routledge</publisher-name>.</citation></ref>
<ref id="ref36"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Halliday</surname> <given-names>S. E.</given-names></name> <name><surname>Calkins</surname> <given-names>S. D.</given-names></name> <name><surname>Leerkes</surname> <given-names>E. M.</given-names></name></person-group> (<year>2018</year>). <article-title>Measuring preschool learning engagement in the laboratory</article-title>. <source>J. Exp. Child Psychol.</source> <volume>167</volume>, <fpage>93</fpage>&#x2013;<lpage>116</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.jecp.2017.10.006</pub-id>, PMID: <pub-id pub-id-type="pmid">29154033</pub-id></citation></ref>
<ref id="ref37"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>He</surname> <given-names>L.</given-names></name> <name><surname>Min</surname> <given-names>S.</given-names></name></person-group> (<year>2017</year>). <article-title>Development and validation of a computer adaptive EFL test</article-title>. <source>Lang. Assess. Q.</source> <volume>14</volume>, <fpage>160</fpage>&#x2013;<lpage>176</lpage>. doi: <pub-id pub-id-type="doi">10.1080/15434303.2016.1162793</pub-id></citation></ref>
<ref id="ref38"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hindman</surname> <given-names>A. H.</given-names></name> <name><surname>Morrison</surname> <given-names>F. J.</given-names></name> <name><surname>Connor</surname> <given-names>C. M.</given-names></name> <name><surname>Connor</surname> <given-names>J. A.</given-names></name></person-group> (<year>2020</year>). <article-title>Bringing the science of reading to preservice elementary teachers: tools that bridge research and practice</article-title>. <source>Read. Res. Q.</source> <volume>55</volume>, <fpage>S197</fpage>&#x2013;<lpage>S206</lpage>. doi: <pub-id pub-id-type="doi">10.1002/rrq.345</pub-id></citation></ref>
<ref id="ref39"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Hirsh-Pasek</surname> <given-names>K.</given-names></name> <name><surname>Zosh</surname> <given-names>J. M.</given-names></name> <name><surname>Golinkoff</surname> <given-names>R. M.</given-names></name> <name><surname>Gray</surname> <given-names>J. H.</given-names></name> <name><surname>Robb</surname> <given-names>M. B.</given-names></name> <name><surname>Kaufman</surname> <given-names>J.</given-names></name></person-group> (<year>2015</year>). <article-title>Putting education in &#x201C;educational&#x201D; apps: lessons from the science of learning</article-title>. <source>Psychol. Sci. Public Interest</source> <volume>16</volume>, <fpage>3</fpage>&#x2013;<lpage>34</lpage>. doi: <pub-id pub-id-type="doi">10.1177/1529100615569721</pub-id>, PMID: <pub-id pub-id-type="pmid">25985468</pub-id></citation></ref>
<ref id="ref40"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ho</surname> <given-names>J. C. S.</given-names></name> <name><surname>McBride</surname> <given-names>C.</given-names></name> <name><surname>Lui</surname> <given-names>K. F. H.</given-names></name> <name><surname>&#x0141;ockiewicz</surname> <given-names>M.</given-names></name></person-group> (<year>2024</year>). <article-title>WordSword: an efficient online word Reading assessment for global English</article-title>. <source>Assessment</source> <volume>31</volume>, <fpage>875</fpage>&#x2013;<lpage>891</lpage>. doi: <pub-id pub-id-type="doi">10.1177/10731911231194971</pub-id>, PMID: <pub-id pub-id-type="pmid">37658620</pub-id></citation></ref>
<ref id="ref41"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Jim&#x00E9;nez</surname> <given-names>M. V.</given-names></name> <name><surname>Yumus</surname> <given-names>M.</given-names></name> <name><surname>Schiele</surname> <given-names>T.</given-names></name> <name><surname>Mues</surname> <given-names>A.</given-names></name> <name><surname>Niklas</surname> <given-names>F.</given-names></name></person-group> (<year>2024</year>). <article-title>Preschool emergent literacy skills as predictors of reading and spelling in grade 2 and the role of migration background in Germany</article-title>. <source>J. Exp. Child Psychol.</source> <volume>244</volume>:<fpage>105927</fpage>. doi: <pub-id pub-id-type="doi">10.1016/j.jecp.2024.105927</pub-id>, PMID: <pub-id pub-id-type="pmid">38678807</pub-id></citation></ref>
<ref id="ref42"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Justice</surname> <given-names>L. M.</given-names></name> <name><surname>Ezell</surname> <given-names>H. K.</given-names></name></person-group> (<year>2001</year>). <article-title>Word and print awareness in 4-year-old children</article-title>. <source>Child Lang. Teach. Ther.</source> <volume>17</volume>, <fpage>207</fpage>&#x2013;<lpage>225</lpage>. doi: <pub-id pub-id-type="doi">10.1177/026565900101700303</pub-id></citation></ref>
<ref id="ref43"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Justice</surname> <given-names>L. M.</given-names></name> <name><surname>Invernizzi</surname> <given-names>M. A.</given-names></name> <name><surname>Meier</surname> <given-names>J. D.</given-names></name></person-group> (<year>2002</year>). <article-title>Designing and implementing an early literacy screening protocol</article-title>. <source>Lang. Speech Hear. Serv. Sch.</source> <volume>33</volume>, <fpage>84</fpage>&#x2013;<lpage>101</lpage>. doi: <pub-id pub-id-type="doi">10.1044/0161-1461(2002/007)</pub-id></citation></ref>
<ref id="ref44"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Justice</surname> <given-names>L. M.</given-names></name> <name><surname>Kaderavek</surname> <given-names>J. N.</given-names></name> <name><surname>Fan</surname> <given-names>X.</given-names></name> <name><surname>Sofka</surname> <given-names>A.</given-names></name> <name><surname>Hunt</surname> <given-names>A.</given-names></name></person-group> (<year>2009</year>). <article-title>Accelerating preschoolers' early literacy development through classroom-based teacher&#x2013;child storybook reading and explicit print referencing</article-title>. <source>Lang. Speech Hear. Serv. Sch.</source> <volume>40</volume>, <fpage>67</fpage>&#x2013;<lpage>85</lpage>. doi: <pub-id pub-id-type="doi">10.1044/0161-1461(2008/07-0098)</pub-id>, PMID: <pub-id pub-id-type="pmid">19124650</pub-id></citation></ref>
<ref id="ref45"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Keuning</surname> <given-names>J.</given-names></name> <name><surname>Verhoeven</surname> <given-names>L.</given-names></name></person-group> (<year>2008</year>). <article-title>Spelling development throughout the elementary grades: the Dutch case</article-title>. <source>Learn. Individ. Differ.</source> <volume>18</volume>, <fpage>459</fpage>&#x2013;<lpage>470</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.lindif.2007.12.001</pub-id></citation></ref>
<ref id="ref46"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Lee</surname> <given-names>Y. H.</given-names></name> <name><surname>Jia</surname> <given-names>Y.</given-names></name></person-group> (<year>2014</year>). <article-title>Using response time to investigate students' test-taking behaviors in a NAEP computer-based study</article-title>. <source>Large-Scale Assess. Educ.</source> <volume>2</volume>, <fpage>1</fpage>&#x2013;<lpage>24</lpage>. doi: <pub-id pub-id-type="doi">10.1186/s40536-014-0008-1</pub-id>, PMID: <pub-id pub-id-type="pmid">40022721</pub-id></citation></ref>
<ref id="ref002"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Levine</surname> <given-names>D. A.</given-names></name> <name><surname>Duncan</surname> <given-names>P. W.</given-names></name> <name><surname>Nguyen-Huynh</surname> <given-names>M. N.</given-names></name> <name><surname>Ogedegbe</surname> <given-names>O. G.</given-names></name></person-group> (<year>2020</year>). <article-title>Interventions targeting racial/ethnic disparities in stroke prevention and treatment</article-title>. <source>Stroke.</source> doi: <pub-id pub-id-type="doi">10.1161/STROKEAHA.120.030427</pub-id></citation></ref>
<ref id="ref004"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liao</surname> <given-names>Z. Y.</given-names></name> <name><surname>Jian</surname> <given-names>F.</given-names></name> <name><surname>Long</surname> <given-names>H.</given-names></name> <name><surname>Lu</surname> <given-names>Y.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Yang</surname> <given-names>Z.</given-names></name> <etal/></person-group>. (<year>2012</year>). <article-title>Validity assessment and determination of the cutoff value for the Index of Complexity, Outcome and Need among 12&#x2013;13 year-olds in Southern Chinese</article-title>. <source>Int J Oral Sci</source> <volume>4</volume>, <fpage>88</fpage>&#x2013;<lpage>93</lpage>. doi: <pub-id pub-id-type="doi">10.1038/ijos.2012.24</pub-id></citation></ref>
<ref id="ref47"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Magis</surname> <given-names>D.</given-names></name> <name><surname>Barrada</surname> <given-names>J. R.</given-names></name></person-group> (<year>2017</year>). <article-title>Computerized adaptive testing with R: recent updates of the package catR</article-title>. <source>J. Stat. Softw.</source> <volume>76</volume>, <fpage>1</fpage>&#x2013;<lpage>19</lpage>. doi: <pub-id pub-id-type="doi">10.18637/jss.v076.c01</pub-id></citation></ref>
<ref id="ref48"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Magis</surname> <given-names>D.</given-names></name> <name><surname>Ra&#x00EE;che</surname> <given-names>G.</given-names></name></person-group> (<year>2012</year>). <article-title>Random generation of response patterns under computerized adaptive testing with the R package catR</article-title>. <source>J. Stat. Softw.</source> <volume>48</volume>, <fpage>1</fpage>&#x2013;<lpage>31</lpage>. doi: <pub-id pub-id-type="doi">10.18637/jss.v048.i08</pub-id></citation></ref>
<ref id="ref49"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Marsh</surname> <given-names>J.</given-names></name> <name><surname>Plowman</surname> <given-names>L.</given-names></name> <name><surname>Yamada-Rice</surname> <given-names>D.</given-names></name> <name><surname>Bishop</surname> <given-names>J.</given-names></name> <name><surname>Scott</surname> <given-names>F.</given-names></name></person-group> (<year>2020</year>). &#x201C;<article-title>Digital play: a new classification</article-title>&#x201D; in eds. <person-group person-group-type="editor"><name><surname>Stephen</surname> <given-names>C.</given-names></name> <name><surname>Edwards</surname> <given-names>S.</given-names></name></person-group>.  <source>Digital play and Technologies in the Early Years</source> (<publisher-loc>London</publisher-loc>: <publisher-name>Routledge</publisher-name>), <fpage>20</fpage>&#x2013;<lpage>31</lpage>.</citation></ref>
<ref id="ref50"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Maydeu-Olivares</surname> <given-names>A.</given-names></name></person-group> (<year>2013</year>). <article-title>Goodness-of-fit assessment of item response theory models</article-title>. <source>Meas. Interdiscip. Res. Perspect.</source> <volume>11</volume>, <fpage>71</fpage>&#x2013;<lpage>101</lpage>. doi: <pub-id pub-id-type="doi">10.1080/15366367.2013.831680</pub-id>, PMID: <pub-id pub-id-type="pmid">39989647</pub-id></citation></ref>
<ref id="ref51"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>McArthur</surname> <given-names>G. M.</given-names></name> <name><surname>Filardi</surname> <given-names>N.</given-names></name> <name><surname>Francis</surname> <given-names>D. A.</given-names></name> <name><surname>Boyes</surname> <given-names>M. E.</given-names></name> <name><surname>Badcock</surname> <given-names>N. A.</given-names></name></person-group> (<year>2020</year>). <article-title>Self-concept in poor readers: a systematic review and meta-analysis</article-title>. <source>PeerJ</source> <volume>8</volume>:<fpage>e8772</fpage>. doi: <pub-id pub-id-type="doi">10.7717/peerj.8772</pub-id>, PMID: <pub-id pub-id-type="pmid">32211239</pub-id></citation></ref>
<ref id="ref52"><citation citation-type="book"><person-group person-group-type="editor"><name><surname>McElvany</surname> <given-names>N.</given-names></name> <name><surname>Lorenz</surname> <given-names>R.</given-names></name> <name><surname>Frey</surname> <given-names>A.</given-names></name> <name><surname>Goldhammer</surname> <given-names>F.</given-names></name> <name><surname>Schilcher</surname> <given-names>A.</given-names></name> <name><surname>Stubbe</surname> <given-names>T. C.</given-names></name></person-group> (Eds.) (<year>2023</year>). <source>IGLU 2021: Lesekompetenz von Grundschulkindern im internationalen Vergleich und im Trend &#x00FC;ber 20 Jahre</source>. <publisher-loc>M&#x00FC;nster, Germany</publisher-loc>: <publisher-name>Waxmann Verlag</publisher-name>.</citation></ref>
<ref id="ref53"><citation citation-type="book"><person-group person-group-type="author"><name><surname>McLachlan</surname> <given-names>C.</given-names></name> <name><surname>Fleer</surname> <given-names>M.</given-names></name> <name><surname>Edwards</surname> <given-names>S.</given-names></name></person-group> (<year>2018</year>). <source>Early childhood curriculum: Planning, assessment and implementation</source>. <publisher-loc>Cambridge</publisher-loc>: <publisher-name>Cambridge University Press</publisher-name>.</citation></ref>
<ref id="ref54"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Meijer</surname> <given-names>R. R.</given-names></name> <name><surname>Nering</surname> <given-names>M. L.</given-names></name></person-group> (<year>1999</year>). <article-title>Computerized adaptive testing: overview and introduction</article-title>. <source>Appl. Psychol. Meas.</source> <volume>23</volume>, <fpage>187</fpage>&#x2013;<lpage>194</lpage>. doi: <pub-id pub-id-type="doi">10.1177/01466219922031310</pub-id></citation></ref>
<ref id="ref55"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Meindl</surname> <given-names>M.</given-names></name> <name><surname>Jungmann</surname> <given-names>T.</given-names></name></person-group> (<year>2019</year>). <source>Erz&#x00E4;hl- und Lesekompetenzen erfassen bei vier- bis f&#x00FC;nfj&#x00E4;hrigen Kindern (EuLe 4&#x2013;5)</source>. <publisher-loc>G&#x00F6;ttingen</publisher-loc>: <publisher-name>Hogrefe</publisher-name>.</citation></ref>
<ref id="ref56"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Morrow</surname> <given-names>L. M.</given-names></name></person-group> (<year>2007</year>). <source>Developing literacy in preschool</source>. <publisher-loc>New York</publisher-loc>: <publisher-name>The Guilford Press</publisher-name>.</citation></ref>
<ref id="ref57"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Nagy</surname> <given-names>W. E.</given-names></name> <name><surname>Anderson</surname> <given-names>R. C.</given-names></name></person-group> (<year>1995</year>) <italic>Metalinguistic awareness and literacy acquisition in different languages</italic>. Center for the Study of Reading Technical Report; no. 618.</citation></ref>
<ref id="ref58"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Nelson</surname> <given-names>P. M.</given-names></name> <name><surname>Van Norman</surname> <given-names>E. R.</given-names></name> <name><surname>Klingbeil</surname> <given-names>D. A.</given-names></name> <name><surname>Parker</surname> <given-names>D. C.</given-names></name></person-group> (<year>2017</year>). <article-title>Progress monitoring with computer adaptive assessments: the impact of data collection schedule on growth estimates</article-title>. <source>Psychol. Sch.</source> <volume>54</volume>, <fpage>463</fpage>&#x2013;<lpage>471</lpage>. doi: <pub-id pub-id-type="doi">10.1002/pits.22015</pub-id></citation></ref>
<ref id="ref59"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Neumann</surname> <given-names>M. M.</given-names></name></person-group> (<year>2018</year>). <article-title>Using tablets and apps to enhance emergent literacy skills in young children</article-title>. <source>Early Child. Res. Q.</source> <volume>42</volume>, <fpage>239</fpage>&#x2013;<lpage>246</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.ecresq.2017.10.006</pub-id></citation></ref>
<ref id="ref60"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Neumann</surname> <given-names>M. M.</given-names></name> <name><surname>Anthony</surname> <given-names>J. L.</given-names></name> <name><surname>Erazo</surname> <given-names>N. A.</given-names></name> <name><surname>Neumann</surname> <given-names>D. L.</given-names></name></person-group> (<year>2019</year>) &#x2018;<article-title>Assessment and technology: mapping future directions in the early childhood classroom</article-title>&#x2019;, in <source>Front. Educ.</source> <volume>4</volume>, p. <fpage>116</fpage>. doi: <pub-id pub-id-type="doi">10.3389/feduc.2019.00116</pub-id></citation></ref>
<ref id="ref61"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Neumann</surname> <given-names>M. M.</given-names></name> <name><surname>Neumann</surname> <given-names>D. L.</given-names></name></person-group> (<year>2014</year>). <article-title>Touch screen tablets and emergent literacy</article-title>. <source>Early Childhood Educ. J.</source> <volume>42</volume>, <fpage>231</fpage>&#x2013;<lpage>239</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10643-013-0608-3</pub-id></citation></ref>
<ref id="ref62"><citation citation-type="other"><person-group person-group-type="author"><collab id="coll2">OECD</collab></person-group> (<year>2013</year>) <italic>Synergies for better learning: An international perspective on evaluation and assessment</italic>, OECD Reviews of Evaluation and Assessment in Education. Available online at: <ext-link xlink:href="https://dx.doi.org/10.1787/9789264190658-en" ext-link-type="uri">https://dx.doi.org/10.1787/9789264190658-en</ext-link> (accessed October 30, 2024).</citation></ref>
<ref id="ref63"><citation citation-type="other"><person-group person-group-type="author"><collab id="coll3">OECD</collab></person-group>. (<year>2015</year>) <italic>Students, computers and learning: Making the connection</italic>. PISA. Available online at: <ext-link xlink:href="https://doi.org/10.1787/9789264239555-en" ext-link-type="uri">https://doi.org/10.1787/9789264239555-en</ext-link> (accessed October 30, 2024)</citation></ref>
<ref id="ref64"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Petermann</surname> <given-names>F.</given-names></name></person-group> (<year>2018</year>). <source>SET 5&#x2013;10. Sprachstandserhebungstest f&#x00FC;r Kinder im Alter zwischen 5 und 10 Jahren</source>. <publisher-loc>G&#x00F6;ttingen</publisher-loc>: <publisher-name>Hogrefe</publisher-name>.</citation></ref>
<ref id="ref65"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Petermann</surname> <given-names>F.</given-names></name> <name><surname>Ri&#x00DF;ling</surname> <given-names>J.-K.</given-names></name> <name><surname>Metzer</surname> <given-names>J.</given-names></name></person-group> (<year>2016</year>). <source>SET 3&#x2013;5. Sprachstandserhebungstest f&#x00FC;r Kinder im Alter zwischen 3 und 5 Jahren, vol. 48</source>. <publisher-loc>G&#x00F6;ttingen</publisher-loc>: <publisher-name>Hogrefe</publisher-name>, <fpage>69</fpage>&#x2013;<lpage>79</lpage>.</citation></ref>
<ref id="ref66"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Piasta</surname> <given-names>S. B.</given-names></name> <name><surname>Justice</surname> <given-names>L. M.</given-names></name> <name><surname>McGinty</surname> <given-names>A. S.</given-names></name> <name><surname>Kaderavek</surname> <given-names>J. N.</given-names></name></person-group> (<year>2012</year>). <article-title>Increasing young children&#x2019;s contact with print during shared reading: longitudinal effects on literacy achievement</article-title>. <source>Child Dev.</source> <volume>83</volume>, <fpage>810</fpage>&#x2013;<lpage>820</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1467-8624.2012.01754.x</pub-id>, PMID: <pub-id pub-id-type="pmid">22506889</pub-id></citation></ref>
<ref id="ref67"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Raymond</surname> <given-names>M. R.</given-names></name> <name><surname>Stevens</surname> <given-names>C.</given-names></name> <name><surname>Bucak</surname> <given-names>S. D.</given-names></name></person-group> (<year>2019</year>). <article-title>The optimal number of options for multiple-choice questions on high-stakes tests: application of a revised index for detecting nonfunctional distractors</article-title>. <source>Adv. Health Sci. Educ.</source> <volume>24</volume>, <fpage>141</fpage>&#x2013;<lpage>150</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10459-018-9855-9</pub-id>, PMID: <pub-id pub-id-type="pmid">30362027</pub-id></citation></ref>
<ref id="ref68"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Roberts</surname> <given-names>J. S.</given-names></name> <name><surname>Donoghue</surname> <given-names>J. R.</given-names></name> <name><surname>Laughlin</surname> <given-names>J. E.</given-names></name></person-group> (<year>2000</year>). <article-title>A general item response theory model for unfolding unidimensional polytomous responses</article-title>. <source>Appl. Psychol. Meas.</source> <volume>24</volume>, <fpage>3</fpage>&#x2013;<lpage>32</lpage>. doi: <pub-id pub-id-type="doi">10.1177/01466216000241001</pub-id></citation></ref>
<ref id="ref69"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Robitzsch</surname> <given-names>A.</given-names></name></person-group> (<year>2022</year>). <article-title>Four-parameter guessing model and related item response models</article-title>. <source>Math. Comput. Appl.</source> <volume>27</volume>:<fpage>95</fpage>. doi: <pub-id pub-id-type="doi">10.3390/mca27060095</pub-id></citation></ref>
<ref id="ref70"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schildkamp</surname> <given-names>K.</given-names></name> <name><surname>Kuiper</surname> <given-names>W.</given-names></name></person-group> (<year>2010</year>). <article-title>Data-informed curriculum reform: which data, what purposes, and promoting and hindering factors</article-title>. <source>Teach. Teach. Educ.</source> <volume>26</volume>, <fpage>482</fpage>&#x2013;<lpage>496</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.tate.2009.06.007</pub-id></citation></ref>
<ref id="ref71"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Segall</surname> <given-names>D. O.</given-names></name></person-group> (<year>2009</year>). &#x201C;<article-title>Principles of multidimensional adaptive testing</article-title>&#x201D; in eds. <person-group person-group-type="editor"><name><surname>van der Linden</surname> <given-names>W. J.</given-names></name> <name><surname>Glas</surname> <given-names>C. A. W.</given-names></name></person-group>.  <source>Elements of adaptive testing</source> (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>Springer</publisher-name>), <fpage>57</fpage>&#x2013;<lpage>75</lpage>.</citation></ref>
<ref id="ref72"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sharkey</surname> <given-names>N. S.</given-names></name> <name><surname>Murnane</surname> <given-names>R. J.</given-names></name></person-group> (<year>2006</year>). <article-title>Tough choices in designing a formative assessment system</article-title>. <source>Am. J. Educ.</source> <volume>112</volume>, <fpage>572</fpage>&#x2013;<lpage>588</lpage>. doi: <pub-id pub-id-type="doi">10.1086/505060</pub-id></citation></ref>
<ref id="ref73"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Shaywitz</surname> <given-names>S. E.</given-names></name></person-group> (<year>1998</year>). <article-title>Dyslexia</article-title>. <source>N. Engl. J. Med.</source> <volume>338</volume>, <fpage>307</fpage>&#x2013;<lpage>312</lpage>. doi: <pub-id pub-id-type="doi">10.1056/NEJM199801293380507</pub-id></citation></ref>
<ref id="ref74"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Sideridis</surname> <given-names>G. D.</given-names></name> <name><surname>Tsaousis</surname> <given-names>I.</given-names></name> <name><surname>Al-Sadaawi</surname> <given-names>A.</given-names></name></person-group> (<year>2018</year>). <article-title>Assessing construct validity in math achievement: an application of multilevel structural equation modeling (MSEM)</article-title>. <source>Front. Psychol.</source> <volume>9</volume>:<fpage>1451</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2018.01451</pub-id>, PMID: <pub-id pub-id-type="pmid">30233437</pub-id></citation></ref>
<ref id="ref75"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Snow</surname> <given-names>C. E.</given-names></name></person-group> (<year>2020</year>). <source>The science of early literacy development: Insights from research on reading acquisition</source>. <publisher-loc>Cambridge, MA</publisher-loc>: <publisher-name>Harvard University Press</publisher-name>.</citation></ref>
<ref id="ref76"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Snow</surname> <given-names>C. E.</given-names></name> <name><surname>Dickinson</surname> <given-names>D. K.</given-names></name></person-group> (<year>1991</year>). &#x201C;<article-title>Skills that aren't basic in a new conception of literacy</article-title>&#x201D; in eds. <person-group person-group-type="editor"><name><surname>Pearson</surname> <given-names>P. D.</given-names></name> <name><surname>Barr</surname> <given-names>R.</given-names></name> <name><surname>Kamil</surname> <given-names>M.</given-names></name></person-group>. <source>The handbook of literacy research</source> (<publisher-loc>New York, NY</publisher-loc>: <publisher-name>Routledge</publisher-name>), <fpage>200</fpage>&#x2013;<lpage>222</lpage>.</citation></ref>
<ref id="ref77"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Snow</surname> <given-names>C. E.</given-names></name> <name><surname>Van Hemel</surname> <given-names>S. B.</given-names></name></person-group> (<year>2008</year>). <source>Early childhood assessment: Why, what, and how</source>. <publisher-loc>Washington, DC</publisher-loc>: <publisher-name>The National Academies Press</publisher-name>.</citation></ref>
<ref id="ref78"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Templin</surname> <given-names>J.</given-names></name> <name><surname>Henson</surname> <given-names>R. A.</given-names></name></person-group> (<year>2010</year>). <source>Diagnostic measurement: Theory, methods, and applications</source>. <publisher-loc>New York, NY</publisher-loc>: <publisher-name>Guilford Press</publisher-name>.</citation></ref>
<ref id="ref79"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tomasik</surname> <given-names>M. J.</given-names></name> <name><surname>Berger</surname> <given-names>S.</given-names></name> <name><surname>Moser</surname> <given-names>U.</given-names></name></person-group> (<year>2018</year>). <article-title>On the development of a computer-based tool for formative student assessment: epistemological, methodological, and practical issues</article-title>. <source>Front. Psychol.</source> <volume>9</volume>:<fpage>2245</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fpsyg.2018.02245</pub-id>, PMID: <pub-id pub-id-type="pmid">30515119</pub-id></citation></ref>
<ref id="ref001"><citation citation-type="book"><person-group person-group-type="author"><name><surname>Vygotsky</surname> <given-names>L. S.</given-names></name></person-group> (<year>1978</year>). <article-title>Mind in society: The development of higher psychological processes</article-title>. <publisher-name>Harvard University Press</publisher-name>. vol. <volume>86</volume>.</citation></ref>
<ref id="ref80"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Weiss</surname> <given-names>D. J.</given-names></name></person-group> (<year>1982</year>). <article-title>Improving measurement quality and efficiency with adaptive testing</article-title>. <source>Appl. Psychol. Meas.</source> <volume>6</volume>, <fpage>473</fpage>&#x2013;<lpage>492</lpage>. doi: <pub-id pub-id-type="doi">10.1177/014662168200600408</pub-id></citation></ref>
<ref id="ref81"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Weiss</surname> <given-names>D. J.</given-names></name></person-group> (<year>2004</year>). <article-title>Computerized adaptive testing for effective and efficient measurement in counseling and education</article-title>. <source>Meas. Eval. Couns. Dev.</source> <volume>37</volume>, <fpage>70</fpage>&#x2013;<lpage>84</lpage>. doi: <pub-id pub-id-type="doi">10.1080/07481756.2004.11909751</pub-id></citation></ref>
<ref id="ref82"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Weiss</surname> <given-names>D. J.</given-names></name> <name><surname>Kingsbury</surname> <given-names>G. G.</given-names></name></person-group> (<year>1984</year>). <article-title>Application of computerized adaptive testing to educational problems</article-title>. <source>J. Educ. Meas.</source> <volume>21</volume>, <fpage>361</fpage>&#x2013;<lpage>375</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1745-3984.1984.tb01040.x</pub-id></citation></ref>
<ref id="ref008"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Whitehurst</surname> <given-names>G. J.</given-names></name> <name><surname>Lonigan</surname> <given-names>C. J.</given-names></name></person-group> (<year>1998</year>). <article-title>Child development and emergent literacy</article-title>. <source>Child Dev.</source> <volume>69</volume>, <fpage>848</fpage>&#x2013;<lpage>872</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.14678624.1998.tb06247.x</pub-id></citation></ref>
<ref id="ref83"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wise</surname> <given-names>S. L.</given-names></name></person-group> (<year>2014</year>). <article-title>The utility of adaptive testing in addressing the problem of unmotivated examinees</article-title>. <source>J. Comput. Adapt. Test.</source> <volume>2</volume>, <fpage>1</fpage>&#x2013;<lpage>17</lpage>. doi: <pub-id pub-id-type="doi">10.7333/jcat.v2i0.30</pub-id></citation></ref>
<ref id="ref84"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wise</surname> <given-names>S. L.</given-names></name> <name><surname>Kong</surname> <given-names>X.</given-names></name></person-group> (<year>2005</year>). <article-title>Response time effort: a new measure of examinee motivation in computer-based tests</article-title>. <source>Appl. Meas. Educ.</source> <volume>18</volume>, <fpage>163</fpage>&#x2013;<lpage>183</lpage>. doi: <pub-id pub-id-type="doi">10.1207/s15324818ame1802_2</pub-id></citation></ref>
<ref id="ref85"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wolf</surname> <given-names>E. J.</given-names></name> <name><surname>Harrington</surname> <given-names>K. M.</given-names></name> <name><surname>Clark</surname> <given-names>S. L.</given-names></name> <name><surname>Miller</surname> <given-names>M. W.</given-names></name></person-group> (<year>2013</year>). <article-title>Sample size requirements for structural equation models: an evaluation of power, bias, and solution propriety</article-title>. <source>Educ. Psychol. Meas.</source> <volume>73</volume>, <fpage>913</fpage>&#x2013;<lpage>934</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0013164413495237</pub-id>, PMID: <pub-id pub-id-type="pmid">25705052</pub-id></citation></ref>
<ref id="ref86"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yao</surname> <given-names>L.</given-names></name></person-group> (<year>2013</year>). <article-title>Comparing the performance of five multidimensional CAT selection procedures with different stopping rules</article-title>. <source>Appl. Psychol. Meas.</source> <volume>37</volume>, <fpage>3</fpage>&#x2013;<lpage>23</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0146621612455687</pub-id></citation></ref>
<ref id="ref87"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Yen</surname> <given-names>Y. C.</given-names></name> <name><surname>Ho</surname> <given-names>R. G.</given-names></name> <name><surname>Laio</surname> <given-names>W. W.</given-names></name> <name><surname>Chen</surname> <given-names>L. J.</given-names></name> <name><surname>Kuo</surname> <given-names>C. C.</given-names></name></person-group> (<year>2012</year>). <article-title>An empirical evaluation of the slip correction in the four parameter logistic models with computerized adaptive testing</article-title>. <source>Appl. Psychol. Meas.</source> <volume>36</volume>, <fpage>75</fpage>&#x2013;<lpage>87</lpage>. doi: <pub-id pub-id-type="doi">10.1177/0146621611432862</pub-id></citation></ref>
</ref-list>
</back>
</article>