<?xml version="1.0" encoding="utf-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Neurosci.</journal-id>
<journal-title>Frontiers in Neuroscience</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Neurosci.</abbrev-journal-title>
<issn pub-type="epub">1662-453X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fnins.2025.1643554</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Neuroscience</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Time-frequency feature calculation of multi-stage audiovisual neural processing via electroencephalogram microstates</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name><surname>Xi</surname> <given-names>Yang</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<xref ref-type="corresp" rid="c001"><sup>&#x002A;</sup></xref>
<uri xlink:href="https://loop.frontiersin.org/people/795183/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Zhang</surname> <given-names>Lu</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Li</surname> <given-names>Cunzhen</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Lv</surname> <given-names>Xiaopeng</given-names></name>
<xref ref-type="aff" rid="aff2"><sup>2</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
</contrib>
<contrib contrib-type="author">
<name><surname>Lan</surname> <given-names>Zhu</given-names></name>
<xref ref-type="aff" rid="aff1"><sup>1</sup></xref>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
</contrib>
</contrib-group>
<aff id="aff1"><sup>1</sup><institution>School of Computer Science, Northeast Electric Power University</institution>, <addr-line>Jilin</addr-line>, <country>China</country></aff>
<aff id="aff2"><sup>2</sup><institution>Department of Chemoradiotherapy, Jilin City Hospital of Chemical Industry</institution>, <addr-line>Jilin</addr-line>, <country>China</country></aff>
<author-notes>
<fn fn-type="edited-by" id="fn0001">
<p>Edited by: Celia Andreu-S&#x00E1;nchez, Autonomous University of Barcelona, Spain</p>
</fn>
<fn fn-type="edited-by" id="fn0002">
<p>Reviewed by: Jos&#x00E9; M. Delgado-Garc&#x00ED;a, Universidad Pablo de Olavide, Spain</p>
<p>Madiha Rehman, Khwaja Fareed University of Engineering and Information Technology (KFUEIT), Pakistan</p>
</fn>
<corresp id="c001">&#x002A;Correspondence: Yang Xi, <email>474465389@qq.com</email></corresp>
</author-notes>
<pub-date pub-type="epub">
<day>07</day>
<month>08</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>19</volume>
<elocation-id>1643554</elocation-id>
<history>
<date date-type="received">
<day>10</day>
<month>06</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>22</day>
<month>07</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#x00A9; 2025 Xi, Zhang, Li, Lv and Lan.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Xi, Zhang, Li, Lv and Lan</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<sec id="sec1001">
<title>Introduction</title>
<p>Audiovisual (AV) perception is a fundamental modality for environmental cognition and social communication, involving complex, non-linear multisensory processing of large-scale neuronal activity modulated by attention. However, precise characterization of the underlying AV processing dynamics remains elusive.</p>
</sec>
<sec id="sec2001">
<title>Methods</title>
<p>We designed an AV semantic discrimination task to acquire electroencephalogram (EEG) data under attended and unattended conditions. To temporally resolve the neural processing stages, we developed an EEG microstate-based analysis method. This involved segmenting the EEG into functional sub-stages by applying hierarchical clustering to global field power-peak topographic maps. The optimal number of microstate classes was determined using the Krzanowski-Lai criterion and Global Explained Variance evaluation. We analyzed filtered EEG data across frequency bands to quantify microstate attributes (e.g., duration, occurrence, coverage, transition probabilities), deriving comprehensive time-frequency features. These features were then used to classify processing states with multiple machine learning models.</p>
</sec>
<sec id="sec3001">
<title>Results</title>
<p>Distinct, temporally continuous microstate sequences were identified characterizing attended versus unattended AV processing. The analysis of microstate attributes yielded time-frequency features that achieved high classification accuracy: 97.8% for distinguishing attended vs. unattended states and 98.6% for discriminating unimodal (auditory or visual) versus multimodal (AV) processing across the employed machine learning models.</p>
</sec>
<sec id="sec4001">
<title>Discussion</title>
<p>Our EEG microstate-based method effectively characterizes the spatio-temporal dynamics of AV processing. Furthermore, it provides neurophysiologically interpretable explanations for the highly accurate classification outcomes, offering significant insights into the neural mechanisms underlying attended and unattended multisensory integration.</p>
</sec>
</abstract>
<kwd-group>
<kwd>audiovisual processing</kwd>
<kwd>electroencephalography</kwd>
<kwd>microstates</kwd>
<kwd>time-frequency features</kwd>
<kwd>attentional mechanism</kwd>
</kwd-group>
<counts>
<fig-count count="10"/>
<table-count count="7"/>
<equation-count count="3"/>
<ref-count count="32"/>
<page-count count="20"/>
<word-count count="10097"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Brain Imaging Methods</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec sec-type="intro" id="sec1">
<label>1</label>
<title>Introduction</title>
<p>One of the fundamental scientific issues in the new era of artificial intelligence is the brain-like cognition of multisensory data, among which the machine understanding of images and sounds is a crucial component (<xref ref-type="bibr" rid="ref11">Khaleghi et al., 2013</xref>). Although the information processing capabilities of computers have improved rapidly, their structure and methods of information processing differ significantly from those of the human brain, resulting in significant differences in cognitive abilities compared to the human brain (<xref ref-type="bibr" rid="ref11">Khaleghi et al., 2013</xref>).</p>
<p>Vision and hearing are the primary means through which humans acquire external information. Studying the neural mechanisms underlying human audiovisual (AV) information processing and analyzing the associated brain activity characteristics can provide a theoretical foundation for developing advanced brain-inspired cognitive algorithms. Such algorithms, designed based on neural principles, hold the potential to significantly enhance computers&#x2019; perception and understanding of the real world. However, research on task-state EEG microstates, especially in AV integration, remains preliminary, and their spatiotemporal dynamics have not yet been systematically elucidated. Current studies predominantly employ resting-state microstate classification methods to analyze task-state data &#x2013; an approach that may lead to feature extraction biases, as task-state EEG involves the coordinated activation of specific neural circuits, whose microstate characteristics may differ significantly from resting-state patterns (<xref ref-type="bibr" rid="ref17">Liu et al., 2020</xref>; <xref ref-type="bibr" rid="ref6">D&#x2019;Croz-Baron et al., 2021</xref>). Further, the microstate sequences involved in AV integration may contain superimposed elements such as attentional modulation and multisensory fusion, which existing methods struggle to interpret effectively. This limitation in feature interpretability prevents researchers from precisely mapping specific microstates to different processing stages of AV integration, thereby hindering research progress in parsing multisensory information integration mechanisms at the level of neural oscillations (<xref ref-type="bibr" rid="ref9">Kaiser et al., 2021</xref>).</p>
<p>When humans perceive the external environment, information received simultaneously through two sensory channels from the same spatial location is often perceived as originating from the same object or event, making that object easier to detect compared to unimodal sensory input (<xref ref-type="bibr" rid="ref27">Starke et al., 2017</xref>; <xref ref-type="bibr" rid="ref16">Li et al., 2017</xref>). The brain effectively merges information from AV sensory channels into a unified, coherent, and robust perceptual process, known as AV integration (<xref ref-type="bibr" rid="ref27">Starke et al., 2017</xref>; <xref ref-type="bibr" rid="ref10">Keil and Senkowski, 2018</xref>). Researchers utilize EEG technology, which offers high temporal resolution and non-invasive advantages, to study the neural processes underlying the processing of AV information by the brain. By analyzing the timing of event-related potentials (ERPs) induced by AV stimuli, the temporal progression of the processing of such information by the brain can be described. Currently, the AV integration process is conventionally partitioned into early perceptual and late cognitive processing stages. For example, ERP components observed around 40&#x202F;ms post-stimulus are classified as early perceptual processing (<xref ref-type="bibr" rid="ref21">Molholm et al., 2002</xref>), whereas those detected at 420&#x2013;580&#x202F;ms are associated with late cognitive processing (<xref ref-type="bibr" rid="ref21">Molholm et al., 2002</xref>). However, many ERP components identified in studies cannot be simply categorized into either of these two stages based on timing alone. For instance, Talsma and Woldorff, using grating patterns and pure tone pulses as AV stimuli, identified ERP components occurring around 190&#x202F;ms, 250&#x202F;ms, and 350&#x2013;450&#x202F;ms after stimulus presentation (<xref ref-type="bibr" rid="ref28">Talsma and Woldorff, 2005</xref>). Since the division between early and late stages is relative and lacks a clear boundary, these ERP components cannot be directly classified into either stage (<xref ref-type="bibr" rid="ref2">Botelho et al., 2023</xref>; <xref ref-type="bibr" rid="ref25">Sala et al., 2025</xref>). Recent studies emphasize that ERP temporal windows are not strictly fixed but dynamically modulated by task demands and cognitive contexts, leading to overlapping or ambiguous component classifications (<xref ref-type="bibr" rid="ref2">Botelho et al., 2023</xref>).</p>
<p>The EEG studies reveal that, during both early and late stages, the topographic neural maps of AV information processing continuously evolve over time. These topographic maps represent the distribution of scalp electric fields, which are closely related to the processing of cognitive tasks by the brain. In other words, the information processing by the brain can be described as a series of alternating, finite types of scalp electric field distributions, known as microstates (<xref ref-type="bibr" rid="ref17">Liu et al., 2020</xref>; <xref ref-type="bibr" rid="ref26">Schiller et al., 2020</xref>; <xref ref-type="bibr" rid="ref24">Ricci et al., 2020</xref>). Microstates are transient representations of the global functional states of the brain, associated with large-scale neural synchronization, and their characteristics reflect the neural activity patterns underlying the processing of current cognitive tasks by the brain (<xref ref-type="bibr" rid="ref17">Liu et al., 2020</xref>). Microstates are defined by the topological structure of the multi-channel electrode recordings from the scalp, remaining stable for a period before rapidly transitioning to other microstates (<xref ref-type="bibr" rid="ref26">Schiller et al., 2020</xref>; <xref ref-type="bibr" rid="ref24">Ricci et al., 2020</xref>).</p>
<p>Combining EEG studies on AV information processing, the electrical activity of the brain during such processing can also be described as a sequence of alternating microstates. Microstates with identical characteristics represent the same brain activity state during AV integration. Therefore, the same type of microstate can represent a sub-stage of AV processing, allowing the activity of the brain during AV information processing to be divided into multiple distinct stages. However, current EEG microstate research primarily focuses on the resting state of the brain, where the study participants receive no external stimuli or tasks and remain in a relaxed state with eyes open or closed (<xref ref-type="bibr" rid="ref19">Michel et al., 2024</xref>). Resting-state EEG microstate studies can explore the differences in brain activity between neurological patients and healthy individuals, providing clinical insights into the neural mechanisms of diseases. Resting-state EEG microstates are generally considered to consist of four types, each representing different cognitive states of the brain (<xref ref-type="bibr" rid="ref13">Koenig and Lehmann, 1996</xref>). Unlike resting-state EEG microstates, the microstates associated with different cognitive tasks are inherently different. Therefore, it is inappropriate to simply apply the four canonical resting-state microstate classes to define task-state EEG activity during AV information processing. Consequently, determining the number of EEG microstates during the processing of AV information tasks by the brain remains a challenge.</p>
<p>Further, the processing of AV information by the brain is an extremely complex process, constantly influenced by attention mechanisms (<xref ref-type="bibr" rid="ref27">Starke et al., 2017</xref>; <xref ref-type="bibr" rid="ref16">Li et al., 2017</xref>). At every moment, the brain receives a vast amount of information from the surrounding environment through different sensory channels. Attention mechanisms enable the brain to continuously make selections from this overwhelming influx of information, influencing the processing of cognitive tasks. ERP studies by <xref ref-type="bibr" rid="ref28">Talsma and Woldorff (2005)</xref> demonstrate that attention plays a crucial role in AV processing, modulating the processing of AV information as early as approximately 80&#x202F;ms after stimulus presentation. <xref ref-type="bibr" rid="ref29">Tang et al. (2016)</xref> proposed an interactive model of attention and AV integration, suggesting that attention can effectively enhance AV integration. Other studies have argued that even in the absence of attention mechanisms, incoming information is still processed by the brain, but the processing mechanisms differ from those when attention is engaged (<xref ref-type="bibr" rid="ref31">Xi et al., 2020a</xref>,<xref ref-type="bibr" rid="ref32">b</xref>; <xref ref-type="bibr" rid="ref7">Fairhall and Macaluso, 2009</xref>; <xref ref-type="bibr" rid="ref28">Talsma and Woldorff, 2005</xref>; <xref ref-type="bibr" rid="ref15">Li et al., 2010</xref>).</p>
<p>The influence of attention on AV information processing is complex, and to date, distinguishing between brain activities under attended and unattended conditions remains challenging. Since EEG signals record the electrical activity generated by large-scale oscillations of neuronal populations in the brain, they contain not only temporal information but also rich frequency-domain features. Similar challenges in decoding perceptive neural responses are observed in visual EEG analysis. For instance (<xref ref-type="bibr" rid="ref22">Rehman et al., 2024</xref>), highlighted that the non-stationarity of EEG signals and high noise levels severely limit classification accuracy in rapid-event visual tasks, achieving only 33.17% accuracy for 40 object classes despite advanced deep learning fusion techniques. Neural oscillations in different frequency bands are closely related to the cognitive states of the brain, which carry distinct &#x201C;meanings&#x201D; and &#x201C;functions.&#x201D; (<xref ref-type="bibr" rid="ref10">Keil and Senkowski, 2018</xref>; <xref ref-type="bibr" rid="ref23">Ren et al., 2020</xref>; <xref ref-type="bibr" rid="ref30">Wang et al., 2018</xref>). For example, the oscillatory responses in the delta, theta, alpha, and beta bands are involved in sensory processing and play roles in different stages and regions (<xref ref-type="bibr" rid="ref10">Keil and Senkowski, 2018</xref>). Delta band oscillations are associated with attentional-selective and processing-attended task-relevant events (<xref ref-type="bibr" rid="ref5">Clayton et al., 2015</xref>). The theta band is predominantly implicated in cognitive control and short-term memory in AV integration (<xref ref-type="bibr" rid="ref10">Keil and Senkowski, 2018</xref>; <xref ref-type="bibr" rid="ref23">Ren et al., 2020</xref>) and plays an important role in bimodal AV stimulus processing (<xref ref-type="bibr" rid="ref23">Ren et al., 2020</xref>). The alpha band is primarily associated with sensory-information maintenance, cognitive control, and distractor suppression (<xref ref-type="bibr" rid="ref1">Bonnefond and Jensen, 2012</xref>), while the beta band activity is associated with decision-making and motor responses for a current task (<xref ref-type="bibr" rid="ref30">Wang et al., 2018</xref>). Given the proposition of <xref ref-type="bibr" rid="ref14">Kopell et al. (2000)</xref> that particular frequency bands of cortical oscillations differentially participate in various cognitive functions, we hypothesized that EEG microstates of different frequency bands are associated with various cognitive functions in the AV processing underlying attention modulation.</p>
<p>In this study, we designed a semantic discrimination experiment involving AV stimuli to acquire EEG data from the brain, processing AV information under both attended (where participants actively directed attention to stimuli) and unattended (where participants ignored stimuli) conditions (<xref ref-type="fig" rid="fig1">Figure 1</xref>). Attended processing engages top-down cognitive control for enhanced sensory integration, while unattended processing relies on automatic bottom-up mechanisms with reduced contextual modulation (<xref ref-type="bibr" rid="ref29">Tang et al., 2016</xref>; <xref ref-type="bibr" rid="ref31">Xi et al., 2020a</xref>,<xref ref-type="bibr" rid="ref32">b</xref>). After preprocessing, we performed hierarchical clustering on the global field power (GFP) peak topographic maps of the EEG data and proposed a Krzanowski-Lai Global Explained Variance (KL_GEV)-based evaluation as an optimal clustering evaluation method. The KL_GEV method integrates the advantages of the KL criterion and GEV, serving as an effective approach for selecting the optimal cluster number in task-state EEG microstate analysis. Its core principle involves identifying the point of diminishing marginal returns in GEV improvement as the cluster number increases to determine the optimal solution. We ultimately obtained microstates for AV information processing under both attended and unattended conditions. Since the same type of microstate represents the same brain activity state, we used the results of microstate clustering to divide the AV information processing into six sub-stages under attended conditions and four sub-stages under unattended conditions. We also calculated and analyzed the properties of these microstates (Duration, Occurrence, Coverage, and Transition Probability) for attended and unattended conditions, exploring the modulation mechanism of attention to AV processing. Additionally, by filtering the EEG data, we computed the microstates and their properties for AV information processing in the delta, theta, alpha, and beta frequency bands, ultimately obtaining time-frequency domain features for AV information processing under both conditions. We validated these features using multiple classifiers, achieving a classification accuracy of up to 97.8% in distinguishing between attended and unattended AV processing. Using the same method, we extracted time-frequency features for AV brain activities separately and classified unimodal visual, unimodal auditory, and AV brain activities, achieving an accuracy of 98.6%. These results demonstrate that the time-frequency domain features obtained through this method effectively characterize AV information processing activity of the brain. The main contributions of this study are as follows:</p>
<list list-type="order">
<list-item>
<p>We have proposed a KL_GEV-based evaluation method for determining the optimal number of clusters in EEG microstates during AV information processing. By integrating the KL criterion with the GEV metric, this method achieves accurate discrimination of the optimal number of microstate clusters for AV EEG data. Compared to traditional optimal clustering evaluation methods, this approach demonstrates superior performance in terms of the Calinski-Harabasz (CH) index and silhouette coefficient, providing a reliable quantitative basis for clustering microstates in AV information processing.</p>
</list-item>
<list-item>
<p>We have proposed a method for dividing sub-stages of AV information processing based on EEG microstates. By leveraging the clustering results of AV microstates, this method uses the time periods of the same microstate category to represent a sub-stage of AV information processing. Consequently, the AV information processing under attended conditions is divided into six sub-stages, while under unattended conditions, it is divided into four sub-stages. This microstate-based division of EEG sub-stages takes into account changes in the cognitive states of the brain, providing a higher-resolution temporal window and offering a new perspective for a deeper understanding of the mechanisms employed by the brain in processing AV information.</p>
</list-item>
<list-item>
<p>We have proposed a method for calculating time-frequency domain features of brain activity during AV information processing by computing microstate properties across multiple frequency bands. We calculated the duration, occurrence frequency, coverage, and transition probability of microstates under both attended and unattended conditions for unfiltered EEG data, as well as delta, theta, alpha, and beta frequency bands. By comparing and analyzing the differences in these microstate properties under attended and unattended conditions, we explored the modulation role of attention in AV information processing. Using these microstate properties as time-frequency domain features to characterize brain activity during AV information processing, we validated their effectiveness through classification experiments on multiple machine learning models, including support vector machine (SVM) and Random Forest models, achieving high classification accuracy. The method for calculating time-frequency domain features in our study effectively characterizes brain activity during AV information processing and provides interpretability for the classification results of machine learning models from the perspective of neural mechanisms of information processing.</p>
</list-item>
</list>
<fig position="float" id="fig1">
<label>Figure 1</label>
<caption>
<p>Schematic of the experimental design in left sessions and right sessions.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g001.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Diagram illustrating experimental sessions divided into left and right sessions. It shows audiovisual (AV), visual (V), and auditory (A) stimuli presentations.  The presentation durations are 400 milliseconds for AV sessions and 300 milliseconds for V sessions, with an inter-stimulus interval of 750 to 1250 milliseconds. A timeline runs horizontally at the bottom.</alt-text>
</graphic>
</fig>
</sec>
<sec sec-type="materials|methods" id="sec2">
<label>2</label>
<title>Materials and methods</title>
<sec id="sec3">
<label>2.1</label>
<title>Participants</title>
<p>This study recruited healthy participants aged between 16 and 30&#x202F;years from Changchun University of Science and Technology, China. The study protocol was reviewed and approved by the Ethics Committee of Changchun University of Science and Technology (Ethical Approval Number: 201705024). The following inclusion criteria were employed for this study: normal vision and hearing, or corrected to normal standards (e.g., wearing glasses or hearing aids); at least a middle school level of education; habitual use of the right hand for most daily tasks; availability of time and willingness to comply with the study requirements. The exclusion criteria were as follows: individuals under guardianship or residing in care institutions; those taking psychiatric medications or with a history of mental illness; individuals with brain diseases or conditions affecting normal brain function; those who have previously participated in similar experiments.</p>
<p>A total of 23 eligible healthy participants (<italic>n</italic>&#x202F;=&#x202F;23; five men and 17 women; age range: 16&#x2013;26&#x202F;years, mean age 22&#x202F;years; education range: 14&#x2013;18&#x202F;years, mean education 15.57&#x202F;&#x00B1;&#x202F;1.56&#x202F;years) were selected to participate in the experimental study. All participants were right-handed, had normal or corrected-to-normal vision and hearing, and completed the experiments without withdrawal. All eligible participants were invited to the International Joint Research Center for Brain Information and Intelligent Science in Jilin Province to understand the experimental procedures and sign a written informed consent form for the experiments.</p>
</sec>
<sec id="sec4">
<label>2.2</label>
<title>Stimuli and experiments</title>
<p>The experiments were conducted in a dimly lit, soundproof room shielded from electronic devices. Participants sat in a comfortable chair with their head position fixed using a chin rest. After receiving instructions from the experimenter about the task, participants completed several practice trials. Once each participant had achieved an accuracy rate exceeding 80%, they were considered to have understood the task. During the formal experiment, each participant performed eight experimental blocks, with each block consisting of 20 auditory (A) stimuli, 20 visual (V) stimuli, and 20 AV stimuli. Among these, four blocks required attending to the left while ignoring the right (termed left-attended blocks), and the remaining four blocks required attending to the right while ignoring the left (termed right-attended blocks). The left-attended and right-attended blocks were conducted in alternated fashion, with participants taking five-minute breaks between each pair of blocks. V stimuli were presented to the left or right of the central fixation point on the display monitor, approximately 6&#x00B0; from the fixation point, with a presentation duration of 300&#x202F;ms. The distance between the central fixation point and the participant&#x2019;s eyes was 80&#x202F;cm. A stimuli were delivered through speakers positioned on both sides of the display, lasting 400&#x202F;ms. For AV stimuli, auditory and visual components started simultaneously with spatial congruency&#x2014;that is, both presented either on the left or right side, lasting 400&#x202F;ms. The inter-stimulus interval varied randomly between 750&#x202F;ms and 1,250&#x202F;ms (<xref ref-type="fig" rid="fig1">Figure 1</xref>).</p>
<p>In the stimulus presentation, stimuli containing images and/or sounds of living objects were considered standard stimuli, while stimuli containing images and/or sounds of non-living objects were considered target stimuli. The frequency of A, V, and AV target stimuli was 20%. Each type of stimulus [2 (standard and target)&#x202F;&#x00D7;&#x202F;3 (A, V, and AV)] appeared in a pseudo-random sequence with equal probability. The stimuli were presented on either the left or right side of the participant with equal probability, following a pseudo-random sequence. The purpose of using a pseudo-random sequence to control the presentation order of A, V, and AV stimuli is to prevent participants from anticipating stimuli due to perceivable patterns. If a fully random sequence were used, the same type of stimulus might appear consecutively multiple times. Therefore, we predetermined the stimulus order to ensure the overall sequence appears irregular to participants. This design prevents participant anticipation from introducing bias while maintaining experimental randomness. Participants were instructed to minimize blinking and body movements to avoid artifacts caused by head motion. They were asked to focus their gaze on the fixation point at the center of the screen, implicitly attending to the stimuli presented on one side while ignoring those on the other side. When hearing and/or seeing a stimulus of a living object, participants were required to press the left button quickly and accurately; when hearing and/or seeing a stimulus of a non-living object, they were to press the right button quickly and accurately.</p>
<p>Each participant was required to complete eight blocks of experiments. Among these, four blocks required attention to the left side while ignoring the right side, referred to as the left-side blocks; the remaining four blocks required attention to the right side while ignoring the left side, referred to as the right-side blocks. The left-side and right-side blocks were alternated, and participants were given a 5-min rest between each pair of experimental blocks.</p>
</sec>
<sec id="sec5">
<label>2.3</label>
<title>EEG data acquisition and preprocessing</title>
<p>In this study, the SynAmps 2 system was used to record EEG signals through a 64-channel electrode cap following the international 10&#x2013;20 system. The AFz electrode served as the ground, and the left mastoid was used as the reference electrode. Horizontal eye movements were recorded via an electrooculogram (EOG) and employing a pair of electrodes placed on the outer sides of the left and right eyes. Vertical eye movements and blinks were recorded by a pair of electrodes placed approximately 1&#x202F;cm above and below the left eye. The EEG and EOG signals were amplified and filtered using an analog bandpass filter with a range of 0.01 to 100&#x202F;Hz. During data acquisition, the impedance was maintained below 5&#x202F;k&#x03A9;. The raw signals were digitized at a sampling rate of 1,000&#x202F;Hz and stored for offline analysis.</p>
<p>EEGLAB software was used for preprocessing and analyzing the EEG data. The EEG was digitally filtered offline with a 0&#x2013;30&#x202F;Hz band-pass. Eye movement artifacts were corrected using independent component analysis. Independent components corresponding to artifact sources and brain activity were separated through a manual procedure. The EEG data were manually screened for residual artifacts and then recomputed to a mean reference. Subsequently, the individual EEG data related to AV stimuli were segmented into epochs, using a time window from &#x2212;200&#x202F;ms to +800&#x202F;ms relative to the AV stimulus onset times. This epoch duration was selected to capture both pre-stimulus baseline activity (&#x2212;200 to 0&#x202F;ms) and post-stimulus neural responses spanning early sensory processing (e.g., P1/N1 components within 40&#x2013;200&#x202F;ms) to late cognitive stages (e.g., N400/P3 components up to 420&#x2013;800&#x202F;ms). A baseline correction was performed using the time window from &#x2212;200&#x202F;ms to 0&#x202F;ms before stimulus onset. Trials with voltage exceeding &#x00B1;100&#x202F;&#x03BC;V at any electrode location (excluding EOG electrodes) were excluded from the analysis. Responses related to false alarms were also removed.</p>
</sec>
<sec id="sec6">
<label>2.4</label>
<title>Division of sub-stages in AV information processing based on EEG microstates</title>
<p>We constructed microstates for the processing of AV information by the brain under both attended and unattended conditions based on raw EEG data. The sub-stages of AV information processing were then divided according to the microstate clustering results. The process of dividing sub-stages in AV information processing based on EEG microstates is illustrated in <xref ref-type="fig" rid="fig2">Figure 2</xref>.</p>
<fig position="float" id="fig2">
<label>Figure 2</label>
<caption>
<p>Flowchart of sub-stage division for AV information processing based on EEG microstates.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g002.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">A four-panel diagram explains a process involving EEG data: (A) shows global field power calculation with EEG waves and corresponding graph. (B) illustrates microstate clustering using TAAHC with peak selection and potential topographic maps. (C) presents optimal cluster number selection using KL_GEV, featuring two charts showing GEV growth and KL_GEV values with microstates MS1 to MS6. (D) depicts microstates backfitting and stage division, including a graph with data sections labeled with numbers and colored bars representing six stages S1 to S6.</alt-text>
</graphic>
</fig>
<p>First, the global field power (GFP) of the AV EEG data was calculated. The topographic maps at the local peak points of the GFP were computed based on the potential values of each electrode. These topographic maps were then subjected to hierarchical clustering to select the optimal number of microstate clusters. Finally, the AV processing stages were divided into multiple sub-stages based on the microstate clustering results.</p>
<sec id="sec7">
<label>2.4.1</label>
<title>Analysis methods for AV EEG microstates</title>
<p>Global field power, as a reference-independent measurement metric, was used to evaluate the overall electrical activity of brain topographic maps. Since local maxima of the GFP curve typically correspond to stable topological configurations and high signal-to-noise ratios, this study first calculated GFP values at each time point within 0&#x2013;800&#x202F;ms and selected scalp potential maps corresponding to local GFP peaks as the initial clustering input prior to cluster analysis. GFP was derived by computing the sum of squared differences between all electrode potentials and the mean potential, with its value reflecting the spatial consistency of EEG signals. Based on AV EEG data under different frequency band conditions, GFP calculation provided highly representative initial topographic maps for subsequent clustering.</p>
<p>Subsequently, a Topographical Atomize and Agglomerate Hierarchical Clustering-based microstate analysis was performed. First, local GFP peak points were identified, with their potential distributions treated as candidate microstates. Initially, the potential map of each local peak was considered an independent cluster. Cluster similarity was measured using Pearson correlation coefficients, and the most similar clusters were iteratively merged to form microstate templates. This bottom-up process gradually reduced the number of clusters, ultimately generating microstate templates for each candidate cluster number.</p>
<p>The selection of the optimal cluster number requires balancing model complexity and interpretability. Common evaluation methods include GEV that maximizes the cumulative GEV to ensure selected clusters sufficiently explain spatial features of EEG signals; cross-validation (CV) which partitions data into training/validation sets to assess model stability across cluster numbers, selecting the number with minimal generalization error; and KL criterion which identifies inflection points where the rate of change in cluster compactness (e.g., within-cluster similarity) significantly declines.</p>
<p>After determining the optimal cluster number, microstate backfitting and refinement were conducted. During backfitting, the potential map of each time point was matched to all microstate templates via Pearson correlation coefficients, assigning it to the most similar class to form continuous microstate time series. This process was iterated until assignments stabilized. Short-duration noise segments were removed using a window smoothing algorithm, yielding stable microstate sequences reflecting AV processing dynamics. This methodology enabled efficient spatiotemporal modeling and analysis of EEG signals. This workflow is illustrated in <xref ref-type="fig" rid="fig3">Figure 3</xref>.</p>
<fig position="float" id="fig3">
<label>Figure 3</label>
<caption>
<p>Flowchart of time-frequency domain feature calculation for AV EEG based on multi-band microstates.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g003.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Diagram showing the process of analyzing EEG data for attended and unattended audiovisual conditions. EEG signals undergo filtering and are divided into delta, theta, alpha, and beta bands. Microstates for each band are visualized as colored head maps. Below, a comparative analysis is performed on microstate properties&#x2014;duration, occurrence, coverage, and transition probability&#x2014;across conditions.</alt-text>
</graphic>
</fig>
</sec>
<sec id="sec8">
<label>2.4.2</label>
<title>Optimal cluster number selection method based on KL_GEV</title>
<p>In this study, we proposed a method for evaluating the clustering results of task-state EEG microstates. The KL_GEV evaluation metric was applied to AV EEG signals, and the microstate classifications selected under the KL_GEV, KL, and CV evaluation metrics were compared based on the variance ratio criterion (CH Index) and the silhouette coefficient.</p>
<p>The calculation principle of GEV is based on the idea of decomposing the total variance into the portion explained by the model and the portion unexplained by the model. By comparing the sizes of these two portions, the proportion of total variance explained by the model can be determined. Specifically, the formula is as follows:</p>
<disp-formula id="E1">
<mml:math id="M1">
<mml:mi mathvariant="italic">GEV</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>t</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>T</mml:mi>
</mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mo stretchy="true">(</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo stretchy="true">)</mml:mo>
<mml:mo>&#x2212;</mml:mo>
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo stretchy="true">&#x0302;</mml:mo>
</mml:mover>
<mml:msup>
<mml:mrow>
<mml:mo stretchy="true">(</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo stretchy="true">)</mml:mo>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mrow>
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>t</mml:mi>
<mml:mo>=</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>T</mml:mi>
</mml:msubsup>
<mml:mi>y</mml:mi>
<mml:mo stretchy="true">(</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo stretchy="true">)</mml:mo>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo stretchy="true">&#x00AF;</mml:mo>
</mml:mover>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
</mml:math>
</disp-formula>
<p>Where <inline-formula>
<mml:math id="M2">
<mml:mi>T</mml:mi>
</mml:math>
</inline-formula> is the total number of time points, and <inline-formula>
<mml:math id="M3">
<mml:mi>y</mml:mi>
<mml:mo stretchy="true">(</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo stretchy="true">)</mml:mo>
</mml:math>
</inline-formula> is the observed value vector at time point <inline-formula>
<mml:math id="M4">
<mml:mi>t</mml:mi>
</mml:math>
</inline-formula>.</p>
<p>KL-GEV is an improved method based on the KL criterion. It determines the optimal number of microstate classifications by analyzing the rate of change of GEV as the number of classifications increases. Its core idea is to capture the &#x201C;point of diminishing marginal returns&#x201D; in GEV growth, i.e., the inflection point where the improvement in model-explained variance significantly weakens with the addition of more classifications. Its calculation formula is as follows:</p>
<disp-formula id="E2">
<mml:math id="M5">
<mml:msub>
<mml:mtext mathvariant="italic">diff</mml:mtext>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:msub>
<mml:mi mathvariant="italic">GEV</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mi mathvariant="italic">GE</mml:mi>
<mml:msub>
<mml:mi>V</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>+</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:math>
</disp-formula>
<disp-formula id="E3">
<mml:math id="M6">
<mml:mi>K</mml:mi>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mrow>
<mml:mi mathvariant="italic">GEV</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>i</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>=</mml:mo>
<mml:mo>&#x2223;</mml:mo>
<mml:mfrac>
<mml:msub>
<mml:mtext mathvariant="italic">diff</mml:mtext>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mi mathvariant="italic">dif</mml:mi>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>+</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#x2223;</mml:mo>
</mml:math>
</disp-formula>
<p>The larger the <inline-formula>
<mml:math id="M7">
<mml:mi>K</mml:mi>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi mathvariant="italic">GEV</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> value, the smaller the impact of adding one more microstate to the growth of GEV. Therefore, the optimal number of microstate classifications can be considered as the point where <inline-formula>
<mml:math id="M8">
<mml:mi>K</mml:mi>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi mathvariant="italic">GEV</mml:mi>
</mml:msub>
</mml:math>
</inline-formula> is maximized.</p>
<p>The KL-GEV microstate clustering number selection criterion was compared with CV and KL using evaluation metrics such as the CH index and the silhouette coefficient.</p>
<p>The CH index essentially represents the ratio of between-cluster distance to within-cluster distance, and its overall calculation process is similar to that of variance, hence it is also referred to as the variance ratio criterion.</p>
<p>The silhouette coefficient measures the separation between clusters by comparing the similarity of each object to its own cluster with its similarity to objects in other clusters. First, the mean distance (within-cluster distance) between sample point and all other sample points in its cluster is calculated, and then the ratio of this distance to the mean distance (between-cluster distance) between the sample point and all sample points in the nearest other cluster is computed. The silhouette coefficient ranges from &#x2212;1 to 1, with higher values indicating better clustering performance.</p>
</sec>
</sec>
<sec id="sec9">
<label>2.5</label>
<title>Features calculation of AV processing modulated by attention mechanisms</title>
<p>As shown in <xref ref-type="fig" rid="fig3">Figure 3</xref>, we calculated the AV processing microstates under both attended and unattended conditions, both before frequency division and within the delta, theta, alpha, and beta frequency bands. By comparing and analyzing the differences in microstates between attended and unattended conditions, we aimed to explore the regulatory role of attention on brain activity during AV processing. Additionally, the microstate properties in each frequency band were used as time-frequency domain features to characterize AV processing. Machine learning algorithms were employed to classify attended and unattended AV processing, further validating the regulatory mechanisms of attention on such processing. This approach also provided a reference method for identifying characteristic computations of brain activity.</p>
<sec id="sec10">
<label>2.5.1</label>
<title>Property calculation for multi-band microstates</title>
<p>In this study, we filtered the preprocessed EEG data during AV stimulus presentation to obtain AV processing EEG data in four frequency bands: delta, theta, alpha, and beta. Using the aforementioned methods for calculating EEG microstates during brain processing of AV tasks, microstates under both attended and unattended conditions were computed for each of the four frequency bands. Together with the original EEG data before frequency division, we ultimately obtained microstates for 2 (attended and unattended conditions)&#x202F;&#x00D7;&#x202F;5 (original, delta, theta, alpha, and beta bands)&#x202F;=&#x202F;10 types of AV processing, as shown in <xref ref-type="fig" rid="fig3">Figure 3</xref>.</p>
<p>Further calculations were performed to determine the following properties for each type of microstate under these 10 conditions: Duration, Coverage, Occurrence, and Transition Probabilities. Among these, Duration was used to describe the length of time each microstate persists in the EEG signal. Coverage was used to describe the proportion of a specific microstate in the entire EEG signal. Occurrence refers to the number of times a specific microstate appears per unit time. The Transition Probability between microstates refers to the likelihood of transitioning from a specific microstate to another microstate.</p>
</sec>
<sec id="sec11">
<label>2.5.2</label>
<title>Time-frequency domain feature analysis of AV processing modulated by attention mechanisms</title>
<p>This study calculated the microstate properties (Duration, Coverage, Occurrence, and Transition Probabilities) under both attended and unattended conditions for the original EEG, as well as the delta, theta, alpha, and beta frequency bands. These microstate time series, which incorporate frequency-domain features, provide a time-frequency domain feature set for analyzing the AV processing stages and the regulatory role of attention. Based on the original EEG, we computed the microstates of AV processing, thereby dividing the processing stages into multiple sub-stages. Similarly, we were able to obtain the sub-stages of AV processing in the delta, theta, alpha, and beta frequency bands.</p>
<p>To explore the characteristics of the sub-stages of AV processing, we compared the microstate properties derived from the original EEG with those from each frequency band. For the microstate properties of Duration, Coverage, Occurrence, and Transition Probability, a repeated measures analysis of variance (ANOVA) (N Microstates conditions [MS1, MS2, &#x2026; and MSn]) as within-subject factors was conducted separately. Paired-t tests were conducted in the presence of main or interaction effects, and for origin EEG and each frequency, all the microstate properties of Duration, Coverage, Occurrence, and Transition Probability were separately entered into a repeated measures ANOVA [two Attention conditions (attended and unattended)] as within-subject factors, to find significant differences caused by the attention mechanism. All statistical analyses were conducted using IBM SPSS software (version 22, IBM Inc., Armonk, NY) for Windows. Greenhouse&#x2013;Geisser corrections were applied with adjusted degrees of freedom. Effects and correlations were considered significant when <italic>p</italic>&#x202F;&#x003C;&#x202F;0.05.</p>
</sec>
<sec id="sec12">
<label>2.5.3</label>
<title>Classification of AV processing brain activities based on time-frequency domain features</title>
<p>To validate that the time-frequency domain features calculated using multi-band microstates could effectively characterize the regulatory role of attention in AV processing, we applied several classical machine learning algorithms to classify brain activities under attended and unattended conditions based on these features. Additionally, we extended this method to classify brain activities during unimodal visual processing, unimodal auditory processing, and AV information processing, further verifying that the time-frequency domain features derived from multi-band microstates can effectively represent brain activities during different information processing tasks.</p>
<p>For the classification experiments in this study, six classical machine learning models were employed: SVM, Random Forest, Gradient Boosting, k-nearest neighbors (KNN), Logistic Regression, and Linear Discriminant Analysis (LDA). The parameter Settings of these six machine learning models are shown in <xref ref-type="table" rid="tab1">Table 1</xref>. The specific experimental setup included an i7-9750H processor, an NVIDIA GeForce RTX 1660 Ti graphics card, 16&#x202F;GB of memory, and the Windows 10 Pro 64-bit operating system. All experiments were conducted in the MATLAB R2022a environment, using 5-fold cross-validation to evaluate the classification performance of each model. Evaluation metrics included accuracy, precision, recall, and F1 score.</p>
<table-wrap position="float" id="tab1">
<label>Table 1</label>
<caption>
<p>Parameter settings of six machine learning models.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Classifier</th>
<th align="left" valign="top">Parameters</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">SVM</td>
<td align="left" valign="top">kernel&#x202F;=&#x202F;&#x2018;rbf&#x2019;, C&#x202F;=&#x202F;1, gamma&#x202F;=&#x202F;1, random_state&#x202F;=&#x202F;42</td>
</tr>
<tr>
<td align="left" valign="top">Random forest</td>
<td align="left" valign="top">n_estimators&#x202F;=&#x202F;100, max_depth&#x202F;=&#x202F;None</td>
</tr>
<tr>
<td align="left" valign="top">Gradient boosting</td>
<td align="left" valign="top">n_estimators&#x202F;=&#x202F;100, max_depth&#x202F;=&#x202F;3</td>
</tr>
<tr>
<td align="left" valign="top">KNN</td>
<td align="left" valign="top"><italic>k</italic> =&#x202F;5</td>
</tr>
<tr>
<td align="left" valign="top">Logistic regression</td>
<td align="left" valign="top">penalty&#x202F;=&#x202F;&#x2018;l2&#x2019;, C&#x202F;=&#x202F;1</td>
</tr>
<tr>
<td align="left" valign="top">LDA</td>
<td align="left" valign="top">solver&#x202F;=&#x202F;&#x2018;svd&#x2019;, tol&#x202F;=&#x202F;0.0001</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
</sec>
</sec>
<sec id="sec13">
<label>3</label>
<title>Results and analysis</title>
<sec id="sec14">
<label>3.1</label>
<title>Division of sub-stages in AV processing</title>
<sec id="sec15">
<label>3.1.1</label>
<title>Evaluation results of optimal cluster number based on KL_GEV</title>
<p>Current research typically divides microstates during the resting state of the brain into four categories: A, B, C, and D. However, when processing AV information, the brain is in a task state, and the changes in scalp electric field distribution are related to the neural mechanisms of AV information processing. Therefore, it is not appropriate to simply categorize microstates in the same way as in the resting state.</p>
<p>We proposed the use of KL_GEV to evaluate the number of microstate clusters, thereby determining the optimal number of microstates for AV information processing. Our KL_GEV calculation results are shown in <xref ref-type="table" rid="tab2">Table 2</xref>. Under attended conditions, the optimal number of microstates for unfiltered AV information processing was six. For the delta, theta, alpha, and beta frequency bands, the optimal numbers of microstates were six, eight, four, and 11, respectively. Under unattended conditions, the optimal number of microstates for unfiltered AV information processing was four. For the delta, theta, alpha, and beta frequency bands, the optimal numbers of microstates were four, 14, five, and eight, respectively.</p>
<table-wrap position="float" id="tab2">
<label>Table 2</label>
<caption>
<p>Comparison of clustering performance based on different evaluation criteria.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top" colspan="2">Conditions</th>
<th align="left" valign="top">Evaluation methods</th>
<th align="center" valign="top">Optimal Cluster Number</th>
<th align="center" valign="top">CH_Score</th>
<th align="center" valign="top">silhouette coefficient</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle" rowspan="15">Attended AV</td>
<td align="left" valign="middle" rowspan="3">Origin</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">7</td>
<td align="center" valign="middle">2164.352</td>
<td align="center" valign="middle">0.07196</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">5</td>
<td align="center" valign="middle">2378.497</td>
<td align="center" valign="middle">0.076566</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">6</td>
<td align="center" valign="middle"><bold>2447.977</bold></td>
<td align="center" valign="middle"><bold>0.076588</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Delta</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">6</td>
<td align="center" valign="middle">2942.375</td>
<td align="center" valign="middle">0.076338</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">10</td>
<td align="center" valign="middle">2219.615</td>
<td align="center" valign="middle">0.033569</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">6</td>
<td align="center" valign="middle"><bold>2942.375</bold></td>
<td align="center" valign="middle"><bold>0.076338</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Theta</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">9</td>
<td align="center" valign="middle">1760.172</td>
<td align="center" valign="middle">&#x2212;0.002210</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">12</td>
<td align="center" valign="middle">1450.215</td>
<td align="center" valign="middle">&#x2212;0.027301</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">8</td>
<td align="center" valign="middle"><bold>1947.016</bold></td>
<td align="center" valign="middle"><bold>0.013379</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Alpha</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">8</td>
<td align="center" valign="middle">2045.314</td>
<td align="center" valign="middle">0.070799</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">13</td>
<td align="center" valign="middle">1453.455</td>
<td align="center" valign="middle">0.052390</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">4</td>
<td align="center" valign="middle"><bold>3247.124</bold></td>
<td align="center" valign="middle"><bold>0.105336</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Beta</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">11</td>
<td align="center" valign="middle">1108.778</td>
<td align="center" valign="middle">0.050355</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">12</td>
<td align="center" valign="middle">1035.411</td>
<td align="center" valign="middle">0.048540</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">11</td>
<td align="center" valign="middle"><bold>1108.778</bold></td>
<td align="center" valign="middle"><bold>0.050355</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="15">Unattended AV</td>
<td align="left" valign="middle" rowspan="3">Origin</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">5</td>
<td align="center" valign="middle">2726.437</td>
<td align="center" valign="middle">0.061659</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">7</td>
<td align="center" valign="middle">2206.027</td>
<td align="center" valign="middle">0.038566</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">4</td>
<td align="center" valign="middle"><bold>3357.477</bold></td>
<td align="center" valign="middle"><bold>0.076578</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Delta</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">4</td>
<td align="center" valign="middle">4545.957</td>
<td align="center" valign="middle">0.112436</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">4</td>
<td align="center" valign="middle">4545.957</td>
<td align="center" valign="middle">0.112436</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">4</td>
<td align="center" valign="middle"><bold>4545.957</bold></td>
<td align="center" valign="middle"><bold>0.112436</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Theta</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">14</td>
<td align="center" valign="middle"><bold>1483.109</bold></td>
<td align="center" valign="middle"><bold>&#x2212;0.018532</bold></td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">13</td>
<td align="center" valign="middle">1405.437</td>
<td align="center" valign="middle">&#x2212;0.020005</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">14</td>
<td align="center" valign="middle"><bold>1483.109</bold></td>
<td align="center" valign="middle"><bold>&#x2212;0.018532</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Alpha</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">12</td>
<td align="center" valign="middle">1540.616</td>
<td align="center" valign="middle">0.045729</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">12</td>
<td align="center" valign="middle">1540.616</td>
<td align="center" valign="middle">0.045729</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">5</td>
<td align="center" valign="middle"><bold>2704.912</bold></td>
<td align="center" valign="middle"><bold>0.084078</bold></td>
</tr>
<tr>
<td align="left" valign="middle" rowspan="3">Beta</td>
<td align="left" valign="middle">CV</td>
<td align="center" valign="middle">7</td>
<td align="center" valign="middle">1321.498</td>
<td align="center" valign="middle">0.052318</td>
</tr>
<tr>
<td align="left" valign="middle">KL</td>
<td align="center" valign="middle">10</td>
<td align="center" valign="middle">1117.718</td>
<td align="center" valign="middle">0.043486</td>
</tr>
<tr>
<td align="left" valign="middle">KL_GEV</td>
<td align="center" valign="middle">8</td>
<td align="center" valign="middle"><bold>1440.693</bold></td>
<td align="center" valign="middle"><bold>0.062208</bold></td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>CH, Calinski-Harabasz. KL, Krzanowski-Lai. KL_GEV, Krzanowski-Lai Global Explained Variance. CV, cross-validation. The bold values in the table represent the optimal values.</p>
</table-wrap-foot>
</table-wrap>
<p>We employed the CH Score and silhouette coefficient as evaluation metrics for clustering performance, comparing them with two other classical methods for determining the optimal number of clusters: CV and KL.</p>
<p>When selecting the number of microstate categories based on CV, two main approaches were used. The first involved observing the slope changes in the CV curve. Although CV continues to decrease as the number of microstate categories increases, the rate of decrease may significantly slow down after a certain number of categories (the &#x201C;elbow point&#x201D;), which can be considered a candidate for the optimal number of categories. The second approach involved physiological constraints, referencing typical numbers of microstate categories in task states from existing literature (e.g., four to seven categories) to avoid selecting excessively high and less interpretable numbers.</p>
<p>The CH Index essentially represents the ratio of between-cluster distance to within-cluster distance, and its calculation process is similar to that of variance, hence it is also referred to as the variance ratio criterion. The silhouette coefficient measures the separation between clusters by comparing the similarity of each object to its own cluster with its similarity to objects in other clusters. The results are shown in <xref ref-type="table" rid="tab2">Table 2</xref> and <xref ref-type="fig" rid="fig4">Figure 4</xref>. These results showcase the CH index and silhouette coefficient of the selected clusters for each microstate clustering method across each frequency band.</p>
<fig position="float" id="fig4">
<label>Figure 4</label>
<caption>
<p>Evaluation of microstate clustering numbers for unfiltered AV processing under attended and unattended conditions.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g004.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Chart (A) shows a line graph with different markers for Attended_KL, Attended_KLGEV, Unattended_KL, and Unattended_KLGEV values across cluster numbers, varying widely. Chart (B) shows a line graph for Attended_CV and Unattended_CV, both decreasing as cluster numbers increase. Shaded areas indicate different cluster ranges.</alt-text>
</graphic>
</fig>
<p>From the above results, it can be observed that the brain processes AV information regardless of whether attentional resources are engaged. Under attended conditions, the AV information processing can be represented by an alternating sequence of six microstates, while under unattended conditions, it is represented by four microstates. This may suggest that the processing of AV information by the brain is more refined and complex when attentional resources are engaged. This finding aligns with conclusions from some previous studies, further highlighting the importance and value of classifying brain microstates for research purposes. From the perspective of information processing, the handling of AV information by the brain is an extremely complex process involving numerous levels and types of neural activities. By classifying microstates, these complex neural activities can be systematically organized and categorized, clearly revealing the specific patterns of information processing under different conditions. For example, in this study, distinguishing between six microstates under attended conditions and four microstates under unattended conditions allows us to intuitively observe the impact of attentional resource allocation on the refinement and complexity of information processing, providing a framework for a deeper understanding of the information processing mechanisms of the brain. From the perspective of exploring neural mechanisms, different microstates may represent the activation of distinct neural functional modules or neural circuits. Classifying microstates helps us identify the specific neural regions and pathways involved in AV information processing. The alternation of different microstates may reflect the dynamic interactions between these neural regions. By analyzing these microstates, we can better uncover the mysteries of brain neural mechanisms and clarify the specific roles and interrelationships of different neural regions in information processing.</p>
<p>Further, the differing results of microstate clustering across frequency bands for AV information imply that neural oscillations in different frequency bands contribute to the processing of AV information, but the mechanisms vary across bands. By further analyzing the properties of microstate sequences in various frequency bands, we can obtain time-domain and frequency-domain features that characterize brain activity during AV information processing under both attended and unattended conditions.</p>
</sec>
<sec id="sec16">
<label>3.1.2</label>
<title>Results of sub-stage division in attention-modulated AV processing</title>
<p>As shown in <xref ref-type="fig" rid="fig4">Figures 4</xref>, <xref ref-type="fig" rid="fig5">5</xref>, we used the KL_GEV evaluation method to cluster the microstates of AV information processing under attended conditions into six categories and those under unattended conditions into four categories. To facilitate a comparative analysis of the microstate properties under both conditions, we relabeled these microstates based on the similarity of their topographic distributions, as illustrated in <xref ref-type="fig" rid="fig6">Figure 6A</xref>. Many classic studies (<xref ref-type="bibr" rid="ref8">Huang et al., 2018</xref>; <xref ref-type="bibr" rid="ref31">Xi et al., 2020a</xref>,<xref ref-type="bibr" rid="ref32">b</xref>; <xref ref-type="bibr" rid="ref18">Matusz and Eimer, 2013</xref>; <xref ref-type="bibr" rid="ref4">Brunelli&#x00E8;re et al., 2013</xref>) divide the AV information processing stages into early and late phases based on the timing of ERP presentations. However, this division lacks clear temporal boundaries and does not consider whether the scalp electric field distributions are consistent within the same phase. The scalp electric field distribution reflects the neural activity state of the brain during information processing and is closely related to cognitive processes. Therefore, we have proposed that processing stages with identical or similar scalp electric field distributions represent identical or similar cognitive processes.</p>
<fig position="float" id="fig5">
<label>Figure 5</label>
<caption>
<p>Evaluation of microstate clustering numbers for AV processing under attended and unattended conditions across different frequency bands.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g005.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Four sets of graphs labeled A, B, C, and D display data trends. Each set includes two line graphs comparing attended and unattended conditions across a number of clusters, with variations in KL and KLGEV metrics on the left and CV metrics on the right. Shaded regions highlight certain cluster ranges. All graphs use red and blue lines for visual distinction.</alt-text>
</graphic>
</fig>
<fig position="float" id="fig6">
<label>Figure 6</label>
<caption>
<p>Sub-stage division results for AV processing under attended and unattended conditions.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g006.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Diagram showing microstate prototypes under different attention conditions and backfitting results. Part A depicts microstate patterns for attended and unattended conditions, contrasted with four classic resting state microstates. Part B presents graphs of global field power (GFP) over time, with corresponding microstate classifications and time intervals.</alt-text>
</graphic>
</fig>
<p>In this study, based on the clustering results of AV microstates, we used the time series of the same microstate category to represent a sub-stage of AV information processing. Consequently, the AV information processing under attended conditions was divided into six and under unattended conditions into four sub-stages, as shown in <xref ref-type="fig" rid="fig6">Figure 6B</xref>.</p>
<p>Our findings demonstrate that the sub-stages of AV information processing, as delineated by microstate segmentation, do not follow a fixed sequential order but rather operate through dynamic alternation and collaboration to accomplish information processing. This suggests that the processing of external information by the brain involves complex mechanisms that likely encompass multiple cognitive and computational processes. We hypothesized that these sub-stages which are defined by microstates reflect the interactive dynamics of various processing and cognitive mechanisms. Notably, attended AV processing was segmented into six sub-stages (MS1&#x2013;MS6), whereas unattended processing yielded four sub-stages (MS1&#x2013;MS4). This marked difference indicates that the allocation of attentional resources significantly enhances processing complexity, potentially reflecting top-down regulatory mechanisms that facilitate refined integration of multimodal information (e.g., conflict resolution, task switching). The increased number of alternating sub-stages may correspond to a more sophisticated dynamic reorganization of cognitive functions. Even in the absence of attentional engagement, the brain maintains a basic processing of AV information (represented by four microstate clusters), albeit through a simpler mechanism characterized by fewer processing sub-stages. This likely reflects an automatic or passive processing mode that lacks the depth of integration and refinement afforded by attention-guided mechanisms. This reduced sub-stage complexity suggests fundamental differences in neural resource allocation and computational demands between the attended and unattended processing states.</p>
<p>It can be observed that MS1 and MS2 resemble the classical microstate D. Previous research suggests this microstate (particularly associated with the right temporoparietal junction, inferior parietal lobule, and the dorsal attention network) primarily orchestrates attentional resource allocation (<xref ref-type="bibr" rid="ref12">Khanna et al., 2015</xref>). During audiovisual processing, it may participate in integrating visual and auditory information. MS3 and MS4 correspond approximately to microstate C. This microstate (typically linked to core regions of the default mode network, such as the posterior cingulate cortex/precuneus and medial prefrontal cortex) is generally associated with self-referential thinking (e.g., autobiographical memory, introspection) during rest (<xref ref-type="bibr" rid="ref3">Brodbeck et al., 2012</xref>). During audiovisual processing, it may mediate the integration of emotion and perception. MS5 and MS6 are similar to microstate B. Previous studies indicate that microstate B (primarily involving the ventral attention network, including the temporoparietal junction, inferior frontal gyrus, and dorsolateral prefrontal cortex) is mainly associated with visuospatial information processing, attentional shifting, and the monitoring of exogenous stimuli (<xref ref-type="bibr" rid="ref20">Michel and Koenig, 2018</xref>).</p>
</sec>
</sec>
<sec id="sec17">
<label>3.2</label>
<title>Calculation results of multi-band microstate properties</title>
<p>To obtain time-domain and frequency-domain features characterizing AV information processing, we further calculated microstate properties, including Duration, Coverage, Occurrence, and Transition Probability. The calculated microstate properties for AV information processing under attended and unattended conditions are presented in <xref ref-type="table" rid="tab3">Table 3</xref>, while the Transition Probability calculation results are shown in <xref ref-type="fig" rid="fig7">Figure 7</xref>.</p>
<table-wrap position="float" id="tab3">
<label>Table 3</label>
<caption>
<p>EEG microstate properties for AV processing.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Properties</th>
<th align="center" valign="top" colspan="2">Duration (ms)</th>
<th align="center" valign="top" colspan="2">Coverage (%)</th>
<th align="center" valign="top" colspan="2">Occurrence (times)</th>
</tr>
<tr>
<th align="left" valign="top">Microstates</th>
<th align="center" valign="top">Attended</th>
<th align="center" valign="top">Unattended</th>
<th align="center" valign="top">Attended</th>
<th align="center" valign="top">Unattended</th>
<th align="center" valign="top">Attended</th>
<th align="center" valign="top">Unattended</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">MS1</td>
<td align="center" valign="middle">77.60</td>
<td align="center" valign="middle">87.88</td>
<td align="center" valign="middle">15.52</td>
<td align="center" valign="middle">24.41</td>
<td align="center" valign="middle">2.12</td>
<td align="center" valign="middle">2.95</td>
</tr>
<tr>
<td align="left" valign="middle">MS2</td>
<td align="center" valign="middle">70.99</td>
<td align="center" valign="middle">78.03</td>
<td align="center" valign="middle">19.10</td>
<td align="center" valign="middle">24.69</td>
<td align="center" valign="middle">2.66</td>
<td align="center" valign="middle">3.24</td>
</tr>
<tr>
<td align="left" valign="middle">MS3</td>
<td align="center" valign="middle">71.00</td>
<td align="center" valign="middle">90.34</td>
<td align="center" valign="middle">16.60</td>
<td align="center" valign="middle">21.99</td>
<td align="center" valign="middle">2.34</td>
<td align="center" valign="middle">3.47</td>
</tr>
<tr>
<td align="left" valign="middle">MS4</td>
<td align="center" valign="middle">68.56</td>
<td align="center" valign="middle">69.53</td>
<td align="center" valign="middle">15.13</td>
<td align="center" valign="middle">28.90</td>
<td align="center" valign="middle">2.28</td>
<td align="center" valign="middle">2.78</td>
</tr>
<tr>
<td align="left" valign="middle">MS5</td>
<td align="center" valign="middle">61.12</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">15.35</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.28</td>
<td align="center" valign="middle">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">MS6</td>
<td align="center" valign="middle">64.53</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">18.30</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.61</td>
<td align="center" valign="middle">&#x2013;</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>EEG, electroencephalogram. AV, audiovisual.</p>
</table-wrap-foot>
</table-wrap>
<fig position="float" id="fig7">
<label>Figure 7</label>
<caption>
<p>Transition probability matrices for AV microstates under attended and unattended conditions.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g007.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Two heatmaps compare "Attended-AV" and "Unattended-AV" conditions. Each matrix displays numerical values with a color scale from blue (low) to red (high). Circular color-coded brain maps accompany rows and columns, visualizing related data distribution.</alt-text>
</graphic>
</fig>
<p>By filtering the EEG data of AV processing, we further calculated the microstates under both attended and unattended conditions in the delta, theta, alpha, and beta frequency bands. The properties of these microstates are presented in <xref ref-type="table" rid="tab4">Tables 4</xref>, <xref ref-type="table" rid="tab5">5</xref>.</p>
<table-wrap position="float" id="tab4">
<label>Table 4</label>
<caption>
<p>Microstate properties for AV processing under attended conditions across different frequency bands.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top" rowspan="2">Attended AV</th>
<th align="center" valign="top" colspan="4">Duration</th>
<th align="center" valign="top" colspan="4">Coverage</th>
<th align="center" valign="top" colspan="4">Occurrence</th>
</tr>
<tr>
<th align="center" valign="top">Delta</th>
<th align="center" valign="top">Theta</th>
<th align="center" valign="top">Alpha</th>
<th align="center" valign="top">Beta</th>
<th align="center" valign="top">Delta</th>
<th align="center" valign="top">Theta</th>
<th align="center" valign="top">Alpha</th>
<th align="center" valign="top">Beta</th>
<th align="center" valign="top">Delta</th>
<th align="center" valign="top">Theta</th>
<th align="center" valign="top">Alpha</th>
<th align="center" valign="top">Beta</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">MS1</td>
<td align="center" valign="middle">93.02</td>
<td align="center" valign="middle">66.61</td>
<td align="center" valign="middle">55.68</td>
<td align="center" valign="middle">38.17</td>
<td align="center" valign="middle">16.51</td>
<td align="center" valign="middle">16.74</td>
<td align="center" valign="middle">32.93</td>
<td align="center" valign="middle">12.27</td>
<td align="center" valign="middle">1.36</td>
<td align="center" valign="middle">2.61</td>
<td align="center" valign="middle">5.98</td>
<td align="center" valign="middle">2.72</td>
</tr>
<tr>
<td align="left" valign="middle">MS2</td>
<td align="center" valign="middle">91.59</td>
<td align="center" valign="middle">48.32</td>
<td align="center" valign="middle">52.75</td>
<td align="center" valign="middle">43.27</td>
<td align="center" valign="middle">14.35</td>
<td align="center" valign="middle">12.85</td>
<td align="center" valign="middle">28.93</td>
<td align="center" valign="middle">11.60</td>
<td align="center" valign="middle">1.36</td>
<td align="center" valign="middle">2.23</td>
<td align="center" valign="middle">5.49</td>
<td align="center" valign="middle">2.28</td>
</tr>
<tr>
<td align="left" valign="middle">MS3</td>
<td align="center" valign="middle">89.09</td>
<td align="center" valign="middle">50.82</td>
<td align="center" valign="middle">52.49</td>
<td align="center" valign="middle">30.36</td>
<td align="center" valign="middle">12.96</td>
<td align="center" valign="middle">13.87</td>
<td align="center" valign="middle">17.86</td>
<td align="center" valign="middle">7.35</td>
<td align="center" valign="middle">1.14</td>
<td align="center" valign="middle">2.28</td>
<td align="center" valign="middle">3.48</td>
<td align="center" valign="middle">1.63</td>
</tr>
<tr>
<td align="left" valign="middle">MS4</td>
<td align="center" valign="middle">96.46</td>
<td align="center" valign="middle">46.83</td>
<td align="center" valign="middle">48.72</td>
<td align="center" valign="middle">26.98</td>
<td align="center" valign="middle">16.96</td>
<td align="center" valign="middle">11.30</td>
<td align="center" valign="middle">20.28</td>
<td align="center" valign="middle">5.43</td>
<td align="center" valign="middle">1.58</td>
<td align="center" valign="middle">1.79</td>
<td align="center" valign="middle">3.97</td>
<td align="center" valign="middle">1.20</td>
</tr>
<tr>
<td align="left" valign="middle">MS5</td>
<td align="center" valign="middle">131.22</td>
<td align="center" valign="middle">48.41</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">40.06</td>
<td align="center" valign="middle">21.34</td>
<td align="center" valign="middle">9.18</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">9.84</td>
<td align="center" valign="middle">1.52</td>
<td align="center" valign="middle">1.58</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.12</td>
</tr>
<tr>
<td align="left" valign="middle">MS6</td>
<td align="center" valign="middle">112.50</td>
<td align="center" valign="middle">50.27</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">27.25</td>
<td align="center" valign="middle">17.88</td>
<td align="center" valign="middle">11.31</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">6.73</td>
<td align="center" valign="middle">1.36</td>
<td align="center" valign="middle">1.79</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.20</td>
</tr>
<tr>
<td align="left" valign="middle">MS7</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">52.20</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">45.47</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">11.83</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">12.28</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.12</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.45</td>
</tr>
<tr>
<td align="left" valign="middle">MS8</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">56.98</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">37.59</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">12.90</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">7.73</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.01</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.68</td>
</tr>
<tr>
<td align="left" valign="middle">MS9</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">34.62</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">7.44</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.58</td>
</tr>
<tr>
<td align="left" valign="middle">MS10</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">44.17</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">12.52</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.66</td>
</tr>
<tr>
<td align="left" valign="middle">MS11</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">26.41</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">6.81</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.47</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>AV, audiovisual.</p>
</table-wrap-foot>
</table-wrap>
<table-wrap position="float" id="tab5">
<label>Table 5</label>
<caption>
<p>Microstate properties for AV processing under unattended conditions across different frequency bands.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top" rowspan="2">Unattended AV</th>
<th align="center" valign="top" colspan="4">Duration</th>
<th align="center" valign="top" colspan="4">Coverage</th>
<th align="center" valign="top" colspan="4">Occurrence</th>
</tr>
<tr>
<th align="center" valign="top">Delta</th>
<th align="center" valign="top">Theta</th>
<th align="center" valign="top">Alpha</th>
<th align="center" valign="top">Beta</th>
<th align="center" valign="top">Delta</th>
<th align="center" valign="top">Theta</th>
<th align="center" valign="top">Alpha</th>
<th align="center" valign="top">Beta</th>
<th align="center" valign="top">Delta</th>
<th align="center" valign="top">Theta</th>
<th align="center" valign="top">Alpha</th>
<th align="center" valign="top">Beta</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="middle">MS1</td>
<td align="center" valign="middle">140.78</td>
<td align="center" valign="middle">51.41</td>
<td align="center" valign="middle">50.43</td>
<td align="center" valign="middle">48.66</td>
<td align="center" valign="middle">21.59</td>
<td align="center" valign="middle">9.20</td>
<td align="center" valign="middle">25.99</td>
<td align="center" valign="middle">22.76</td>
<td align="center" valign="middle">1.65</td>
<td align="center" valign="middle">1.53</td>
<td align="center" valign="middle">5.11</td>
<td align="center" valign="middle">4.60</td>
</tr>
<tr>
<td align="left" valign="middle">MS2</td>
<td align="center" valign="middle">187.05</td>
<td align="center" valign="middle">45.94</td>
<td align="center" valign="middle">58.62</td>
<td align="center" valign="middle">43.72</td>
<td align="center" valign="middle">32.88</td>
<td align="center" valign="middle">9.02</td>
<td align="center" valign="middle">28.15</td>
<td align="center" valign="middle">11.08</td>
<td align="center" valign="middle">1.99</td>
<td align="center" valign="middle">1.48</td>
<td align="center" valign="middle">4.89</td>
<td align="center" valign="middle">2.44</td>
</tr>
<tr>
<td align="left" valign="middle">MS3</td>
<td align="center" valign="middle">126.72</td>
<td align="center" valign="middle">35.01</td>
<td align="center" valign="middle">47.46</td>
<td align="center" valign="middle">45.66</td>
<td align="center" valign="middle">23.38</td>
<td align="center" valign="middle">8.34</td>
<td align="center" valign="middle">18.41</td>
<td align="center" valign="middle">13.51</td>
<td align="center" valign="middle">1.82</td>
<td align="center" valign="middle">1.42</td>
<td align="center" valign="middle">3.70</td>
<td align="center" valign="middle">2.61</td>
</tr>
<tr>
<td align="left" valign="middle">MS4</td>
<td align="center" valign="middle">116.25</td>
<td align="center" valign="middle">36.95</td>
<td align="center" valign="middle">47.66</td>
<td align="center" valign="middle">35.73</td>
<td align="center" valign="middle">22.15</td>
<td align="center" valign="middle">7.07</td>
<td align="center" valign="middle">15.34</td>
<td align="center" valign="middle">8.73</td>
<td align="center" valign="middle">1.88</td>
<td align="center" valign="middle">1.02</td>
<td align="center" valign="middle">3.10</td>
<td align="center" valign="middle">1.70</td>
</tr>
<tr>
<td align="left" valign="middle">MS5</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">42.89</td>
<td align="center" valign="middle">38.80</td>
<td align="center" valign="middle">46.13</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">6.82</td>
<td align="center" valign="middle">12.11</td>
<td align="center" valign="middle">12.72</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.14</td>
<td align="center" valign="middle">2.61</td>
<td align="center" valign="middle">2.56</td>
</tr>
<tr>
<td align="left" valign="middle">MS6</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">33.33</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">42.59</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">7.63</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">10.62</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.36</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.22</td>
</tr>
<tr>
<td align="left" valign="middle">MS7</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">23.02</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">37.59</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">3.81</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">10.21</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">0.68</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">2.22</td>
</tr>
<tr>
<td align="left" valign="middle">MS8</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">34.45</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">45.19</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">5.77</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">10.38</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.02</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.93</td>
</tr>
<tr>
<td align="left" valign="middle">MS9</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">44.75</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">7.59</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.19</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">MS10</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">57.23</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">10.97</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.59</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">MS11</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">34.30</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">6.45</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.14</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">MS12</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">25.68</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">4.03</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">0.80</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">MS13</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">31.70</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">6.79</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.25</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
</tr>
<tr>
<td align="left" valign="middle">MS14</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">32.73</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">6.51</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">1.19</td>
<td align="center" valign="middle">&#x2013;</td>
<td align="center" valign="middle">&#x2013;</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>AV, audiovisual.</p>
</table-wrap-foot>
</table-wrap>
<p>The Transition Probabilities of microstates across different frequency bands are illustrated in <xref ref-type="fig" rid="fig8">Figure 8</xref>.</p>
<fig position="float" id="fig8">
<label>Figure 8</label>
<caption>
<p>Transition probabilities of microstates for AV processing across different frequency bands.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g008.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Heatmap graphs compare Attended-AV and Unattended-AV conditions at different frequency bands. Each heatmap shows color gradients, with numerical values in grid squares. Surrounding circular maps depict electrophysiological data, with color scales indicating intensity changes. Heatmap with color gradient from blue to red displaying numerical data. Circular inserts with colored patterns are aligned horizontally and vertically. A scale bar ranges from zero to point four.</alt-text>
</graphic>
</fig>
</sec>
<sec id="sec18">
<label>3.3</label>
<title>EEG signal classification results</title>
<p>The classification results based on machine learning models such as SVM, Random Forest, Gradient Boosting, KNN, Logistic Regression, and LDA are shown in <xref ref-type="fig" rid="fig9">Figures 9</xref>, <xref ref-type="fig" rid="fig10">10</xref>. A 5-fold cross-validation was used to evaluate the performance of each classifier. The AV EEG signals were classified into attended and unattended conditions, and the EEG signals under attended conditions were further classified into AV, auditory, and visual categories. The classification results for attended and unattended conditions are shown in <xref ref-type="fig" rid="fig9">Figure 9</xref>.</p>
<fig position="float" id="fig9">
<label>Figure 9</label>
<caption>
<p>Classification results and ROC curves for attended vs. unattended EEG signals.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g009.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">Six machine learning models with confusion matrices and ROC curves. The top row shows confusion matrices for SVM, Random Forest, and Gradient Boosting, indicating high accuracy with most values concentrated on the diagonal. The bottom row displays confusion matrices for KNN, Logistic Regression, and LDA, with similar results. Beneath each matrix are ROC curves for each model, detailing AUC values: SVM (0.902 to 0.935), Random Forest (0.876 to 0.935), Gradient Boosting (0.897 to 0.932), KNN (0.804 to 0.916), Logistic Regression (0.895 to 0.939), and LDA (0.689 to 0.926).</alt-text>
</graphic>
</fig>
<fig position="float" id="fig10">
<label>Figure 10</label>
<caption>
<p>Classification results and ROC curves for AV, auditory, and visual EEG signals.</p>
</caption>
<graphic xlink:href="fnins-19-1643554-g010.tif" mimetype="image" mime-subtype="tiff">
<alt-text content-type="machine-generated">The image contains confusion matrices and ROC curves for six models: SVM, Random Forest, Gradient Boosting, KNN, Logistic Regression, and LDA. Each confusion matrix shows the predictions for three classes: AV, A, and V. ROC curves display performance across five folds, with AUC scores provided for each model. Models generally show high performance, with mean AUCs ranging from 0.78 to 0.94.</alt-text>
</graphic>
</fig>
<p>Additionally, we employed multiple machine learning models to classify AV EEG signals under attended and unattended conditions based on microstate features. The classification results are shown in <xref ref-type="table" rid="tab6">Table 6</xref>, demonstrating that most machine learning models achieve a satisfactory performance in distinguishing between attended and unattended conditions using microstate features.</p>
<table-wrap position="float" id="tab6">
<label>Table 6</label>
<caption>
<p>Classification results for attended and unattended AV processing brain activities based on time-frequency domain features.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Models</th>
<th align="center" valign="top">Accuracy (%)</th>
<th align="center" valign="top">Precision (%)</th>
<th align="center" valign="top">Recall (%)</th>
<th align="center" valign="top">F1-Score (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">SVM</td>
<td align="center" valign="top">97.8</td>
<td align="center" valign="top">98.0</td>
<td align="center" valign="top">90.0</td>
<td align="center" valign="top">97.8</td>
</tr>
<tr>
<td align="left" valign="top">Random forest</td>
<td align="center" valign="top">97.4</td>
<td align="center" valign="top">98.3</td>
<td align="center" valign="top">97.5</td>
<td align="center" valign="top">97.7</td>
</tr>
<tr>
<td align="left" valign="top">Gradient boosting</td>
<td align="center" valign="top">97.4</td>
<td align="center" valign="top">96.3</td>
<td align="center" valign="top">95.5</td>
<td align="center" valign="top">95.4</td>
</tr>
<tr>
<td align="left" valign="top">KNN</td>
<td align="center" valign="top">95.6</td>
<td align="center" valign="top">95.0</td>
<td align="center" valign="top">93.5</td>
<td align="center" valign="top">93.2</td>
</tr>
<tr>
<td align="left" valign="top">Logistic regression</td>
<td align="center" valign="top">97.8</td>
<td align="center" valign="top">98.0</td>
<td align="center" valign="top">98.0</td>
<td align="center" valign="top">97.8</td>
</tr>
<tr>
<td align="left" valign="top">LDA</td>
<td align="center" valign="top">93.3</td>
<td align="center" valign="top">91.2</td>
<td align="center" valign="top">86.0</td>
<td align="center" valign="top">84.9</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>AV, audiovisual. SVM, support vector machine. KNN, k-nearest neighbors. LDA, linear discriminant analysis.</p>
</table-wrap-foot>
</table-wrap>
<p>These results indicate that by dividing the multi-band AV EEG signals into multiple stages using microstates and calculating microstate properties as time-frequency features, most machine learning models can effectively learn and classify the data with an accuracy of approximately 97%. This result also demonstrates that this method can effectively characterize brain activity during the processing of AV information (<xref ref-type="table" rid="tab7">Table 7</xref>).</p>
<table-wrap position="float" id="tab7">
<label>Table 7</label>
<caption>
<p>Classification results for AV, auditory, and visual processing brain activities based on time-frequency domain features.</p>
</caption>
<table frame="hsides" rules="groups">
<thead>
<tr>
<th align="left" valign="top">Models</th>
<th align="center" valign="top">Accuracy (%)</th>
<th align="center" valign="top">Precision (%)</th>
<th align="center" valign="top">Recall (%)</th>
<th align="center" valign="top">F1-Score (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td align="left" valign="top">SVM</td>
<td align="center" valign="top">98.6</td>
<td align="center" valign="top">98.9</td>
<td align="center" valign="top">98.7</td>
<td align="center" valign="top">98.7</td>
</tr>
<tr>
<td align="left" valign="top">Random forest</td>
<td align="center" valign="top">97.1</td>
<td align="center" valign="top">97.8</td>
<td align="center" valign="top">97.3</td>
<td align="center" valign="top">97.3</td>
</tr>
<tr>
<td align="left" valign="top">Gradient boosting</td>
<td align="center" valign="top">98.6</td>
<td align="center" valign="top">98.9</td>
<td align="center" valign="top">98.7</td>
<td align="center" valign="top">98.7</td>
</tr>
<tr>
<td align="left" valign="top">KNN</td>
<td align="center" valign="top">94.3</td>
<td align="center" valign="top">95.7</td>
<td align="center" valign="top">94.3</td>
<td align="center" valign="top">94.2</td>
</tr>
<tr>
<td align="left" valign="top">Logistic regression</td>
<td align="center" valign="top">98.6</td>
<td align="center" valign="top">98.9</td>
<td align="center" valign="top">98.7</td>
<td align="center" valign="top">98.7</td>
</tr>
<tr>
<td align="left" valign="top">LDA</td>
<td align="center" valign="top">69.6</td>
<td align="center" valign="top">66.1</td>
<td align="center" valign="top">69.7</td>
<td align="center" valign="top">65.1</td>
</tr>
</tbody>
</table>
<table-wrap-foot>
<p>AV, audiovisual. SVM, support vector machine. KNN, k-nearest neighbors. LDA, linear discriminant analysis.</p>
</table-wrap-foot>
</table-wrap>
</sec>
</sec>
<sec sec-type="conclusions" id="sec19">
<label>4</label>
<title>Conclusion</title>
<p>This study utilized EEG microstates to divide the AV information processing process into multiple sub-stages and calculated microstate attributes across multiple frequency bands to comprehensively characterize the corresponding brain activity. We propose an evaluation method based on KL_GEV for determining the optimal number of microstate clusters in AV EEG processing, which integrates the KL criterion with the GEV metric to identify the most appropriate number of microstate clusters.</p>
<p>Additionally, this study presented an EEG microstate-based method for segmenting AV information processing into sub-stages. Based on the microstate clustering results, this method used temporally continuous microstate sequences of the same class to represent individual processing sub-stages, thereby dividing attended AV processing into six sub-stages and unattended processing into four sub-stages. This microstate-based segmentation was able to account for changes in cognitive states and provided a higher temporal resolution, offering new perspectives for understanding the neural mechanisms of AV information processing.</p>
<p>Further, by computing microstate attributes across multiple frequency bands, we developed a method for calculating time-frequency domain features of brain activity during AV processing. We calculated the Duration, Occurrence, Coverage, and Transition Probability of microstates in unfiltered data and in delta, theta, alpha, and beta frequency bands under both attended and unattended conditions, comparing the differences in these attributes to investigate the regulatory role of attention in AV processing. Using these frequency-band microstate attributes as time-frequency domain features characterizing AV processing brain activity, we validated their effectiveness through classification with various machine learning models (SVM, Random Forest, etc.). These features achieved up to 97.8% accuracy in classifying attended versus unattended AV processing brain activity and 98.6% accuracy in classifying unimodal (visual, auditory) and multimodal (AV) brain activities. Our time-frequency feature calculation method effectively characterized brain activity during AV information processing and provided neurophysiological interpretability for the machine learning classification results from the perspective of information processing mechanisms. This study provided theoretical and experimental foundations for analyzing the neural mechanisms of multisensory integration and developing brain-inspired information processing algorithms.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="sec20">
<title>Data availability statement</title>
<p>The raw data supporting the conclusions of this article will be made available by the authors, without undue reservation.</p>
</sec>
<sec sec-type="ethics-statement" id="sec21">
<title>Ethics statement</title>
<p>The studies involving humans were approved by Ethics Committee of Changchun University of Science and Technology. The studies were conducted in accordance with the local legislation and institutional requirements. The participants provided their written informed consent to participate in this study.</p>
</sec>
<sec sec-type="author-contributions" id="sec22">
<title>Author contributions</title>
<p>YX: Investigation, Methodology, Supervision, Writing &#x2013; review &#x0026; editing, Writing &#x2013; original draft. LZ: Writing &#x2013; review &#x0026; editing, Methodology, Writing &#x2013; original draft, Data curation, Investigation. CL: Writing &#x2013; review &#x0026; editing, Validation, Investigation, Resources. XL: Writing &#x2013; review &#x0026; editing, Supervision, Validation, Methodology, Resources. ZL: Software, Writing &#x2013; review &#x0026; editing, Data curation.</p>
</sec>
<sec sec-type="funding-information" id="sec23">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research and/or publication of this article. This research was financially supported by the National Natural Science Foundation of China (grant no. 62206044).</p>
</sec>
<sec sec-type="COI-statement" id="sec24">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="ai-statement" id="sec25">
<title>Generative AI statement</title>
<p>The authors declare that no Gen AI was used in the creation of this manuscript.</p>
</sec>
<sec sec-type="disclaimer" id="sec26">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="ref1"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Bonnefond</surname> <given-names>M.</given-names></name> <name><surname>Jensen</surname> <given-names>O.</given-names></name></person-group> (<year>2012</year>). <article-title>Alpha oscillations serve to protect working memory maintenance against anticipated distracters</article-title>. <source>Curr. Biol.</source> <volume>22</volume>, <fpage>1969</fpage>&#x2013;<lpage>1974</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.cub.2012.08.029</pub-id>, PMID: <pub-id pub-id-type="pmid">23041197</pub-id></citation></ref>
<ref id="ref2"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Botelho</surname> <given-names>C.</given-names></name> <name><surname>Fernandes</surname> <given-names>C.</given-names></name> <name><surname>Campos</surname> <given-names>C.</given-names></name> <name><surname>Seixas</surname> <given-names>C.</given-names></name> <name><surname>Pasion</surname> <given-names>R.</given-names></name> <name><surname>Garcez</surname> <given-names>H.</given-names></name> <etal/></person-group>. (<year>2023</year>). <article-title>Uncertainty deconstructed: conceptual analysis and state-of-the-art review of the ERP correlates of risk and ambiguity in decision-making</article-title>. <source>Cogn. Affect. Behav. Neurosci.</source> <volume>23</volume>, <fpage>522</fpage>&#x2013;<lpage>542</lpage>. doi: <pub-id pub-id-type="doi">10.3758/s13415-023-01101-8</pub-id>, PMID: <pub-id pub-id-type="pmid">37173606</pub-id></citation></ref>
<ref id="ref3"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Brodbeck</surname> <given-names>V.</given-names></name> <name><surname>Kuhn</surname> <given-names>A.</given-names></name> <name><surname>von Wegner</surname> <given-names>F.</given-names></name> <name><surname>Morzelewski</surname> <given-names>A.</given-names></name> <name><surname>Tagliazucchi</surname> <given-names>E.</given-names></name> <name><surname>Borisov</surname> <given-names>S.</given-names></name> <etal/></person-group>. (<year>2012</year>). <article-title>EEG microstates of wakefulness and NREM sleep</article-title>. <source>NeuroImage</source> <volume>62</volume>, <fpage>2129</fpage>&#x2013;<lpage>2139</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.neuroimage.2012.05.060</pub-id>, PMID: <pub-id pub-id-type="pmid">22658975</pub-id></citation></ref>
<ref id="ref4"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Brunelli&#x00E8;re</surname> <given-names>A.</given-names></name> <name><surname>S&#x00E1;nchez-Garc&#x00ED;a</surname> <given-names>C.</given-names></name> <name><surname>Ikumi</surname> <given-names>N.</given-names></name> <name><surname>Soto-Faraco</surname> <given-names>S.</given-names></name></person-group> (<year>2013</year>). <article-title>Visual information constrains early and late stages of spoken-word recognition in sentence context</article-title>. <source>Int. J. Psychophysiol.</source> <volume>89</volume>, <fpage>136</fpage>&#x2013;<lpage>147</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.ijpsycho.2013.06.016</pub-id>, PMID: <pub-id pub-id-type="pmid">23797145</pub-id></citation></ref>
<ref id="ref5"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Clayton</surname> <given-names>M. S.</given-names></name> <name><surname>Yeung</surname> <given-names>N.</given-names></name> <name><surname>Kadosh</surname> <given-names>R. C.</given-names></name></person-group> (<year>2015</year>). <article-title>The roles of cortical oscillations in sustained attention</article-title>. <source>Trends Cogn. Sci.</source> <volume>19</volume>, <fpage>188</fpage>&#x2013;<lpage>195</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.tics.2015.02.004</pub-id>, PMID: <pub-id pub-id-type="pmid">25765608</pub-id></citation></ref>
<ref id="ref6"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>D&#x2019;Croz-Baron</surname> <given-names>D. F.</given-names></name> <name><surname>Br&#x00E9;chet</surname> <given-names>L.</given-names></name> <name><surname>Baker</surname> <given-names>M.</given-names></name> <name><surname>Karp</surname> <given-names>T.</given-names></name></person-group> (<year>2021</year>). <article-title>Auditory and visual tasks influence the temporal dynamics of EEG microstates during post-encoding rest</article-title>. <source>Brain Topogr.</source> <volume>34</volume>, <fpage>19</fpage>&#x2013;<lpage>28</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10548-020-00802-4</pub-id>, PMID: <pub-id pub-id-type="pmid">33095401</pub-id></citation></ref>
<ref id="ref7"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Fairhall</surname> <given-names>S. L.</given-names></name> <name><surname>Macaluso</surname> <given-names>E.</given-names></name></person-group> (<year>2009</year>). <article-title>Spatial attention can modulate audiovisual integration at multiple cortical and subcortical sites</article-title>. <source>Eur. J. Neurosci.</source> <volume>29</volume>, <fpage>1247</fpage>&#x2013;<lpage>1257</lpage>. doi: <pub-id pub-id-type="doi">10.1111/j.1460-9568.2009.06688.x</pub-id>, PMID: <pub-id pub-id-type="pmid">19302160</pub-id></citation></ref>
<ref id="ref8"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Huang</surname> <given-names>J.</given-names></name> <name><surname>Reinders</surname> <given-names>A. A. T. S.</given-names></name> <name><surname>Wang</surname> <given-names>Y.</given-names></name> <name><surname>Xu</surname> <given-names>T.</given-names></name> <name><surname>Zeng</surname> <given-names>Y.-w.</given-names></name> <name><surname>Li</surname> <given-names>K.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Neural correlates of audiovisual sensory integration</article-title>. <source>Neuropsychology</source> <volume>32</volume>, <fpage>329</fpage>&#x2013;<lpage>336</lpage>. doi: <pub-id pub-id-type="doi">10.1037/neu0000393</pub-id></citation></ref>
<ref id="ref9"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kaiser</surname> <given-names>M.</given-names></name> <name><surname>Senkowski</surname> <given-names>D.</given-names></name> <name><surname>Keil</surname> <given-names>J.</given-names></name></person-group> (<year>2021</year>). <article-title>Mediofrontal theta-band oscillations reflect top-down influence in the ventriloquist illusion</article-title>. <source>Hum. Brain Mapp.</source> <volume>42</volume>, <fpage>452</fpage>&#x2013;<lpage>466</lpage>. doi: <pub-id pub-id-type="doi">10.1002/hbm.25236</pub-id>, PMID: <pub-id pub-id-type="pmid">33617132</pub-id></citation></ref>
<ref id="ref10"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Keil</surname> <given-names>J.</given-names></name> <name><surname>Senkowski</surname> <given-names>D.</given-names></name></person-group> (<year>2018</year>). <article-title>Neural oscillations orchestrate multisensory processing</article-title>. <source>Neuroscientist</source> <volume>24</volume>, <fpage>609</fpage>&#x2013;<lpage>626</lpage>. doi: <pub-id pub-id-type="doi">10.1177/1073858418755352</pub-id>, PMID: <pub-id pub-id-type="pmid">29424265</pub-id></citation></ref>
<ref id="ref11"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Khaleghi</surname> <given-names>B.</given-names></name> <name><surname>Khamis</surname> <given-names>A.</given-names></name> <name><surname>Karray</surname> <given-names>F. O.</given-names></name> <name><surname>Razavi</surname> <given-names>S. N.</given-names></name></person-group> (<year>2013</year>). <article-title>Multisensor data fusion: a review of the state-of-the-art</article-title>. <source>Informat. Fusion</source> <volume>14</volume>, <fpage>28</fpage>&#x2013;<lpage>44</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.inffus.2011.08.001</pub-id></citation></ref>
<ref id="ref12"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Khanna</surname> <given-names>A.</given-names></name> <name><surname>Pascual-Leone</surname> <given-names>A.</given-names></name> <name><surname>Michel</surname> <given-names>C. M.</given-names></name> <name><surname>Farzan</surname> <given-names>F.</given-names></name></person-group> (<year>2015</year>). <article-title>Microstates in resting-state EEG: current status and future directions</article-title>. <source>Neurosci. Biobehav. Rev.</source> <volume>49</volume>, <fpage>105</fpage>&#x2013;<lpage>113</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.neubiorev.2014.12.010</pub-id>, PMID: <pub-id pub-id-type="pmid">25526823</pub-id></citation></ref>
<ref id="ref13"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Koenig</surname> <given-names>T.</given-names></name> <name><surname>Lehmann</surname> <given-names>D.</given-names></name></person-group> (<year>1996</year>). <article-title>Microstates in language-related brain potential maps show noun&#x2013;verb differences</article-title>. <source>Brain Lang.</source> <volume>53</volume>, <fpage>169</fpage>&#x2013;<lpage>182</lpage>. doi: <pub-id pub-id-type="doi">10.1006/brln.1996.0043</pub-id>, PMID: <pub-id pub-id-type="pmid">8726532</pub-id></citation></ref>
<ref id="ref14"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Kopell</surname> <given-names>N.</given-names></name> <name><surname>Ermentrout</surname> <given-names>G. B.</given-names></name> <name><surname>Whittington</surname> <given-names>M. A.</given-names></name> <name><surname>Traub</surname> <given-names>R. D.</given-names></name></person-group> (<year>2000</year>). <article-title>Gamma rhythms and beta rhythms have different synchronization properties</article-title>. <source>Proc. Natl. Acad. Sci.</source> <volume>97</volume>, <fpage>1867</fpage>&#x2013;<lpage>1872</lpage>. doi: <pub-id pub-id-type="doi">10.1073/pnas.97.4.1867</pub-id>, PMID: <pub-id pub-id-type="pmid">10677548</pub-id></citation></ref>
<ref id="ref15"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>Q.</given-names></name> <name><surname>Wu</surname> <given-names>J.</given-names></name> <name><surname>Touge</surname> <given-names>T.</given-names></name></person-group> (<year>2010</year>). <article-title>Audiovisual interaction enhances auditory detection in late stage: an event-related potential study</article-title>. <source>Neuroreport</source> <volume>21</volume>, <fpage>173</fpage>&#x2013;<lpage>178</lpage>. doi: <pub-id pub-id-type="doi">10.1097/WNR.0b013e3283345f08</pub-id>, PMID: <pub-id pub-id-type="pmid">20065887</pub-id></citation></ref>
<ref id="ref16"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Li</surname> <given-names>Q.</given-names></name> <name><surname>Yu</surname> <given-names>H. T.</given-names></name> <name><surname>Li</surname> <given-names>X. J.</given-names></name> <name><surname>Sun</surname> <given-names>H. Z.</given-names></name> <name><surname>Yang</surname> <given-names>J. J.</given-names></name> <name><surname>Li</surname> <given-names>C. L</given-names></name></person-group>. (<year>2017</year>). <article-title>The informativity of sound modulates crossmodal facilitation of visual discrimination an fMRI study</article-title>. <source>Neuroreport</source> <volume>28</volume>, <fpage>63</fpage>&#x2013;<lpage>68</lpage>.</citation></ref>
<ref id="ref17"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname> <given-names>J.</given-names></name> <name><surname>Xu</surname> <given-names>J.</given-names></name> <name><surname>Zou</surname> <given-names>G.</given-names></name> <name><surname>He</surname> <given-names>Y.</given-names></name> <name><surname>Zou</surname> <given-names>Q.</given-names></name> <name><surname>Gao</surname> <given-names>J. H.</given-names></name></person-group> (<year>2020</year>). <article-title>Reliability and individual specificity of EEG microstate characteristics</article-title>. <source>Brain Topogr.</source> <volume>33</volume>, <fpage>438</fpage>&#x2013;<lpage>449</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10548-020-00777-2</pub-id>, PMID: <pub-id pub-id-type="pmid">32468297</pub-id></citation></ref>
<ref id="ref18"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Matusz</surname> <given-names>P. J.</given-names></name> <name><surname>Eimer</surname> <given-names>M.</given-names></name></person-group> (<year>2013</year>). <article-title>Top-down control of audiovisual search by bimodal search templates</article-title>. <source>Psychophysiology</source> <volume>50</volume>, <fpage>996</fpage>&#x2013;<lpage>1009</lpage>. doi: <pub-id pub-id-type="doi">10.1111/psyp.12086</pub-id>, PMID: <pub-id pub-id-type="pmid">23834379</pub-id></citation></ref>
<ref id="ref19"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Michel</surname> <given-names>C. M.</given-names></name> <name><surname>Brechet</surname> <given-names>L.</given-names></name> <name><surname>Schiller</surname> <given-names>B.</given-names></name> <name><surname>Koenig</surname> <given-names>T.</given-names></name></person-group> (<year>2024</year>). <article-title>Current state of EEG/ERP microstate research</article-title>. <source>Brain Topogr.</source> <volume>37</volume>, <fpage>169</fpage>&#x2013;<lpage>180</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s10548-024-01037-3</pub-id>, PMID: <pub-id pub-id-type="pmid">38349451</pub-id></citation></ref>
<ref id="ref20"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Michel</surname> <given-names>C. M.</given-names></name> <name><surname>Koenig</surname> <given-names>T.</given-names></name></person-group> (<year>2018</year>). <article-title>EEG microstates as a tool for studying the temporal dynamics of whole-brain neuronal networks: a review</article-title>. <source>NeuroImage</source> <volume>180</volume>, <fpage>577</fpage>&#x2013;<lpage>593</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.neuroimage.2017.11.062</pub-id>, PMID: <pub-id pub-id-type="pmid">29196270</pub-id></citation></ref>
<ref id="ref21"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Molholm</surname> <given-names>S.</given-names></name> <name><surname>Ritter</surname> <given-names>W.</given-names></name> <name><surname>Murray</surname> <given-names>M. M.</given-names></name> <name><surname>Javitt</surname> <given-names>D. C.</given-names></name> <name><surname>Schroeder</surname> <given-names>C. E.</given-names></name> <name><surname>Foxe</surname> <given-names>J. J.</given-names></name></person-group> (<year>2002</year>). <article-title>Multisensory auditory&#x2013;visual interactions during early sensory processing in humans: a high-density electrical mapping study</article-title>. <source>Cogn. Brain Res.</source> <volume>14</volume>, <fpage>115</fpage>&#x2013;<lpage>128</lpage>. doi: <pub-id pub-id-type="doi">10.1016/S0926-6410(02)00066-6</pub-id>, PMID: <pub-id pub-id-type="pmid">12063135</pub-id></citation></ref>
<ref id="ref22"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Rehman</surname> <given-names>M.</given-names></name> <name><surname>Anwer</surname> <given-names>H.</given-names></name> <name><surname>Garay</surname> <given-names>H.</given-names></name> <name><surname>Alemany-Iturriaga</surname> <given-names>J.</given-names></name> <name><surname>D&#x00ED;ez</surname> <given-names>I. D. T.</given-names></name> <name><surname>Siddiqui</surname> <given-names>H. R.</given-names></name> <etal/></person-group>. (<year>2024</year>). <article-title>Decoding brain signals from rapid-event EEG for visual analysis using deep learning</article-title>. <source>Sensors</source> <volume>24</volume>:<fpage>6965</fpage>. doi: <pub-id pub-id-type="doi">10.3390/s24216965</pub-id>, PMID: <pub-id pub-id-type="pmid">39517862</pub-id></citation></ref>
<ref id="ref23"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ren</surname> <given-names>Y.</given-names></name> <name><surname>Pan</surname> <given-names>L.</given-names></name> <name><surname>Du</surname> <given-names>X.</given-names></name> <name><surname>Li</surname> <given-names>X.</given-names></name> <name><surname>Hou</surname> <given-names>Y.</given-names></name> <name><surname>Bao</surname> <given-names>J</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>Theta oscillation and functional connectivity alterations related to executive control in temporal lobe epilepsy with comorbid depression</article-title>. <source>Clin. Neurophysiol.</source> <volume>131</volume>, <fpage>1599</fpage>&#x2013;<lpage>1609</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.clinph.2020.03.038</pub-id></citation></ref>
<ref id="ref24"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Ricci</surname> <given-names>L.</given-names></name> <name><surname>Croce</surname> <given-names>P.</given-names></name> <name><surname>Lanzone</surname> <given-names>J.</given-names></name> <name><surname>Boscarino</surname> <given-names>M.</given-names></name> <name><surname>Zappasodi</surname> <given-names>F.</given-names></name> <name><surname>Tombini</surname> <given-names>M.</given-names></name> <etal/></person-group>. (<year>2020</year>). <article-title>Transcutaneous vagus nerve stimulation modulates EEG microstates and delta activity in healthy subjects</article-title>. <source>Brain Sci.</source> <volume>10</volume>:<fpage>668</fpage>. doi: <pub-id pub-id-type="doi">10.3390/brainsci10100668</pub-id>, PMID: <pub-id pub-id-type="pmid">32992726</pub-id></citation></ref>
<ref id="ref25"><citation citation-type="other"><person-group person-group-type="author"><name><surname>Sala</surname> <given-names>M.</given-names></name> <name><surname>Vespignani</surname> <given-names>F.</given-names></name> <name><surname>Gastaldon</surname> <given-names>S.</given-names></name> <name><surname>Casalino</surname> <given-names>L.</given-names></name> <name><surname>Peressotti</surname> <given-names>F</given-names></name></person-group>. (<year>2025</year>). <article-title>ERP evidence of speaker-specific phonological prediction</article-title>.  [Epubh ahead of preprint]. doi: <pub-id pub-id-type="doi">10.1101/2025.04.16.648895</pub-id></citation></ref>
<ref id="ref26"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Schiller</surname> <given-names>B.</given-names></name> <name><surname>Kleinert</surname> <given-names>T.</given-names></name> <name><surname>Teige-Mocigemba</surname> <given-names>S.</given-names></name> <name><surname>Klauer</surname> <given-names>K. C.</given-names></name> <name><surname>Heinrichs</surname> <given-names>M.</given-names></name></person-group> (<year>2020</year>). <article-title>Temporal dynamics of resting EEG networks are associated with prosociality</article-title>. <source>Sci. Rep.</source> <volume>10</volume>:<fpage>13066</fpage>. doi: <pub-id pub-id-type="doi">10.1038/s41598-020-69999-5</pub-id>, PMID: <pub-id pub-id-type="pmid">32747655</pub-id></citation></ref>
<ref id="ref27"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Starke</surname> <given-names>J.</given-names></name> <name><surname>Ball</surname> <given-names>F.</given-names></name> <name><surname>Heinze</surname> <given-names>H. J.</given-names></name> <name><surname>Heinze</surname> <given-names>H.&#x2010;. J.</given-names></name> <name><surname>Noesselt</surname> <given-names>T.</given-names></name></person-group> (<year>2017</year>). <article-title>The spatio-temporal profile of multisensory integration</article-title>. <source>Eur. J. Neurosci.</source> <volume>51</volume>, <fpage>1210</fpage>&#x2013;<lpage>1223</lpage>. doi: <pub-id pub-id-type="doi">10.1111/ejn.13753</pub-id>, PMID: <pub-id pub-id-type="pmid">29057531</pub-id></citation></ref>
<ref id="ref28"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Talsma</surname> <given-names>D.</given-names></name> <name><surname>Woldorff</surname> <given-names>M. G.</given-names></name></person-group> (<year>2005</year>). <article-title>Selective attention and multisensory integration multiple phases of effects on the evoked brain activity</article-title>. <source>J. Cogn. Neurosci.</source> <volume>17</volume>, <fpage>1098</fpage>&#x2013;<lpage>1114</lpage>. doi: <pub-id pub-id-type="doi">10.1162/0898929054475172</pub-id>, PMID: <pub-id pub-id-type="pmid">16102239</pub-id></citation></ref>
<ref id="ref29"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Tang</surname> <given-names>X.</given-names></name> <name><surname>Wu</surname> <given-names>J.</given-names></name> <name><surname>Shen</surname> <given-names>Y.</given-names></name></person-group> (<year>2016</year>). <article-title>The interactions of multisensory integration with endogenous and exogenous attention</article-title>. <source>Neurosci. Biobehav. Rev.</source> <volume>61</volume>, <fpage>208</fpage>&#x2013;<lpage>224</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.neubiorev.2015.11.002</pub-id>, PMID: <pub-id pub-id-type="pmid">26546734</pub-id></citation></ref>
<ref id="ref30"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Wang</surname> <given-names>B.</given-names></name> <name><surname>Li</surname> <given-names>P.</given-names></name> <name><surname>Li</surname> <given-names>D.</given-names></name> <name><surname>Niu</surname> <given-names>Y.</given-names></name> <name><surname>Yan</surname> <given-names>T.</given-names></name> <name><surname>Li</surname> <given-names>T.</given-names></name> <etal/></person-group>. (<year>2018</year>). <article-title>Increased functional brain network efficiency during audiovisual temporal asynchrony integration task in aging</article-title>. <source>Front. Aging Neurosci.</source> <volume>10</volume>:<fpage>316</fpage>. doi: <pub-id pub-id-type="doi">10.3389/fnagi.2018.00316</pub-id>, PMID: <pub-id pub-id-type="pmid">30356825</pub-id></citation></ref>
<ref id="ref31"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xi</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>Q.</given-names></name> <name><surname>Gao</surname> <given-names>N.</given-names></name> <name><surname>Li</surname> <given-names>G.</given-names></name> <name><surname>Lin</surname> <given-names>W.</given-names></name> <name><surname>Wu</surname> <given-names>J.</given-names></name></person-group> (<year>2020a</year>). <article-title>Co-stimulation-removed audiovisual semantic integration and modulation of attention: an event-related potential study</article-title>. <source>Int. J. Psychophysiol.</source> <volume>151</volume>, <fpage>7</fpage>&#x2013;<lpage>17</lpage>. doi: <pub-id pub-id-type="doi">10.1016/j.ijpsycho.2020.02.009</pub-id>, PMID: <pub-id pub-id-type="pmid">32061614</pub-id></citation></ref>
<ref id="ref32"><citation citation-type="journal"><person-group person-group-type="author"><name><surname>Xi</surname> <given-names>Y.</given-names></name> <name><surname>Li</surname> <given-names>Q.</given-names></name> <name><surname>Zhang</surname> <given-names>M.</given-names></name> <name><surname>Liu</surname> <given-names>L.</given-names></name> <name><surname>Wu</surname> <given-names>J.</given-names></name></person-group> (<year>2020b</year>). <article-title>Characterizing the time-varying brain networks of audiovisual integration across frequency bands</article-title>. <source>Cogn. Comput.</source> <volume>12</volume>, <fpage>1154</fpage>&#x2013;<lpage>1169</lpage>. doi: <pub-id pub-id-type="doi">10.1007/s12559-020-09783-9</pub-id></citation></ref>
</ref-list>
</back>
</article>