<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Earth Sci.</journal-id>
<journal-title>Frontiers in Earth Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Earth Sci.</abbrev-journal-title>
<issn pub-type="epub">2296-6463</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">1601090</article-id>
<article-id pub-id-type="doi">10.3389/feart.2025.1601090</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Earth Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>A prototype-based rockburst types and risk prediction algorithm considering intra-class variance and inter-class distance of microseismic data</article-title>
<alt-title alt-title-type="left-running-head">Zhang et al.</alt-title>
<alt-title alt-title-type="right-running-head">
<ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3389/feart.2025.1601090">10.3389/feart.2025.1601090</ext-link>
</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Zhang</surname>
<given-names>Xiufeng</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Li</surname>
<given-names>Guoying</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Chen</surname>
<given-names>Yang</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Wang</surname>
<given-names>Hao</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Zhang</surname>
<given-names>Haikuan</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/2773860/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Li</surname>
<given-names>Haitao</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Du</surname>
<given-names>Weisheng</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1878737/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Li</surname>
<given-names>Xiao</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Xu</surname>
<given-names>Xuewei</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>He</surname>
<given-names>Yuze</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/Writing - review &#x26; editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Coal Industry Management Department</institution>, <institution>Shandong Energy Group Co., Ltd.</institution>, <addr-line>Jinan</addr-line>, <addr-line>Shandong</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>Deep Mining and Rock Burst Research Branch</institution>, <institution>Chinese Institute of Coal Science</institution>, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1639649/overview">Xin Yin</ext-link>, City University of Hong Kong, Hong Kong SAR, China</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1976239/overview">Huan Sun</ext-link>, Hainan University, China</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/3023425/overview">Yunpeng Zhang</ext-link>, China University of Geosciences Wuhan, China</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Haikuan Zhang, <email>820818414@qq.com</email>
</corresp>
</author-notes>
<pub-date pub-type="epub">
<day>29</day>
<month>05</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>13</volume>
<elocation-id>1601090</elocation-id>
<history>
<date date-type="received">
<day>27</day>
<month>03</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>28</day>
<month>04</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2025 Zhang, Li, Chen, Wang, Zhang, Li, Du, Li, Xu and He.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Zhang, Li, Chen, Wang, Zhang, Li, Du, Li, Xu and He</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<p>The prediction and classification of rockburst risk based on microseismic data is the premise of preventing rockbursts during deep mine excavation. By reviewing previous studies, this paper finds two problems that hinder the rockburst prediction: 1) there is a lack of research on the distribution features of monitoring data on the main controlling factors of rockbursts; 2) there is no research on the intra-class variance and inter-class gap of microseismic data. Based on the typical rockburst risk events, a quantitative information model of geology and mining is constructed. The relationship between the spatial&#x2013;temporal distribution characteristics of microseismic data before a rockburst and the main controlling factors of a rockburst is studied. The results show that the distribution features may be different for the same type of microseismic (MS) and rockburst events, and different types of events may show similar distribution features. Therefore, based on the quantitative study of the relationship between the performance of a deep learning prediction algorithm and a rockburst prediction vector, a rockburst risk and type prediction algorithm based on a convolutional neural network (CNN)-gated recurrent unit (GRU) model with prototype-based prediction is proposed. The CNN-GRU model can produce prediction vectors by fusing implicit and explicit information extracted from the original MS data and early warning indicators. Cross-entropy loss, vector-prototype contrastive loss, and vector-prototype contrastive loss are proposed to automatically control the intra-class variance and inter-class gap of prediction vectors belonging to different rockburst risks and types. Many experiments show that the performance of the proposed CNN-GRU model with prototype-based prediction is superior to other algorithms in the prediction of rockburst risks and types based on MS data.</p>
</abstract>
<kwd-group>
<kwd>rockburst prediction</kwd>
<kwd>rockburst types</kwd>
<kwd>deep learning</kwd>
<kwd>microseismic data</kwd>
<kwd>prototype learning</kwd>
</kwd-group>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-at-acceptance</meta-name>
<meta-value>Geohazards and Georisks</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1">
<title>1 Introduction</title>
<p>As underground mining and excavations continue to expand at a rapid pace, a significant engineering challenge arises in the instability of the surrounding rock masses (<xref ref-type="bibr" rid="B4">Aydan et al., 2017</xref>). Rockburst, a representative instability phenomenon, is caused by the abrupt release of accumulated elastic strain energy, posing a grave threat to the safety of workers and causing extensive damage to underground engineering structures (<xref ref-type="bibr" rid="B5">Basnet et al., 2023</xref>). A significant number of researchers are committed to addressing this challenge by elucidating the occurrence mechanism (<xref ref-type="bibr" rid="B3">Askaripour et al., 2022</xref>; <xref ref-type="bibr" rid="B12">He et al., 2023</xref>), enhancing prediction accuracy (<xref ref-type="bibr" rid="B5">Basnet et al., 2023</xref>; <xref ref-type="bibr" rid="B25">Pu et al., 2019</xref>), and developing effective control measures (<xref ref-type="bibr" rid="B18">Li et al., 2019</xref>; <xref ref-type="bibr" rid="B13">He et al., 2018</xref>). Rockburst prediction aims to generate accurate signals prior to the occurrence of these disasters, serving as a prerequisite for controlling and managing the rockburst hazards. Nevertheless, due to its unpredictable emergence and numerous influencing factors, rockburst prediction remains a challenging task that has yet to be fully resolved.</p>
<p>Rockburst prediction is conventionally classified into long-term prediction and short-term prediction (<xref ref-type="bibr" rid="B20">Liang et al., 2020</xref>). Long-term rockburst prediction endeavors to utilize rock mechanical parameters for constructing a prediction model with the aim of assessing the rockburst probability and types of diverse surrounding rock masses under assorted field conditions. This task is customarily accomplished during the initial stage of engineering design or excavation (<xref ref-type="bibr" rid="B22">Liang and Zhao, 2022</xref>). The objective of short-term rockburst prediction, on the other hand, is to predict the time, types, and damage magnitude of dangerous rockburst events by conducting dynamic and static analyses on real-time monitoring data. It is typically carried out during the excavation period (<xref ref-type="bibr" rid="B17">Jinqiang et al., 2021</xref>). To attain satisfactory outcomes, researchers are dedicated to applying diverse methods for rockburst prediction, such as empirical analytical (<xref ref-type="bibr" rid="B13">He et al., 2018</xref>; <xref ref-type="bibr" rid="B28">Yang et al., 2018</xref>), experimental (<xref ref-type="bibr" rid="B14">Hu et al., 2023</xref>; <xref ref-type="bibr" rid="B8">Cheng et al., 2023</xref>), numerical (<xref ref-type="bibr" rid="B26">Wang et al., 2021</xref>; <xref ref-type="bibr" rid="B24">Manouchehrian and Cai, 2018</xref>), intelligent (<xref ref-type="bibr" rid="B1">Adoko and Zvarivadza, 2018</xref>; <xref ref-type="bibr" rid="B27">Xue et al., 2023</xref>), and expert system (<xref ref-type="bibr" rid="B19">Li et al., 2020</xref>) methods. Although each rockburst prediction method has its own advantages, in contrast to intelligent methods, traditional prediction methods hinge on extensive expert experience and meticulous judgment. Therefore, the machine learning method represents a promising alternative and has been adopted by numerous researchers to dissect and handle the intricate and nonlinear process of rockburst prediction.</p>
<p>Microseismic (MS) monitoring is a widely acknowledged and highly effective tool for identifying the dangerous signals and key controlling factors of rockburst types for short-term rockburst prediction. It is capable of monitoring the occurrence of rockbursts by extracting valuable signals that propagate from the fracturing process of rock masses. Recently, by capitalizing on the capabilities of machine learning in handling nonlinear problems, scholars have directed their attention to rockburst prediction based on MS data through the application of existing machine learning algorithms. Such algorithms encompass support vector machine (SVM) (<xref ref-type="bibr" rid="B15">Ji et al., 2020</xref>; <xref ref-type="bibr" rid="B16">Jin et al., 2022</xref>), convolutional neural network (CNN) (<xref ref-type="bibr" rid="B11">Dong et al., 2023</xref>; <xref ref-type="bibr" rid="B34">Zhang et al., 2021</xref>; <xref ref-type="bibr" rid="B30">Yin et al., 2021a</xref>), variants of RNN (<xref ref-type="bibr" rid="B14">Hu et al., 2023</xref>; <xref ref-type="bibr" rid="B9">Di et al., 2023a</xref>; <xref ref-type="bibr" rid="B10">Di et al., 2023b</xref>), convolutional long short-term memory (ConvLSTM) (<xref ref-type="bibr" rid="B6">Chen et al., 2023</xref>; <xref ref-type="bibr" rid="B23">Ma et al., 2021</xref>), and ensemble-learning (<xref ref-type="bibr" rid="B20">Liang et al., 2020</xref>; <xref ref-type="bibr" rid="B32">Yin et al., 2021b</xref>; <xref ref-type="bibr" rid="B21">Liang et al., 2021</xref>), among others (<xref ref-type="bibr" rid="B31">Yin et al., 2024a</xref>; <xref ref-type="bibr" rid="B29">Yin et al., 2024b</xref>; <xref ref-type="bibr" rid="B7">Cheng et al., 2024</xref>; <xref ref-type="bibr" rid="B33">Yin et al., 2021c</xref>). <xref ref-type="bibr" rid="B34">Zhang et al. (2021)</xref> and <xref ref-type="bibr" rid="B30">Yin et al. (2021a)</xref> have conducted an exploration of the key microseismic indexes that can characterize the development process of rockbursts and have applied the refined convolutional neural network (CNN) to predict rockbursts. In order to depict the spatiotemporal relationship within microseismic data and process the spatiotemporal indexes for rockburst prediction, <xref ref-type="bibr" rid="B6">Chen et al. (2023)</xref> developed a deep learning model founded on a ConvLSTM to forecast short-term rockburst risks. Concurrently, the ensemble-learning methodologies (<xref ref-type="bibr" rid="B20">Liang et al., 2020</xref>; <xref ref-type="bibr" rid="B32">Yin et al., 2021b</xref>; <xref ref-type="bibr" rid="B21">Liang et al., 2021</xref>) have also been employed to acquire a highly potent rockburst prediction model relying on MS parameters.</p>
<p>Previous research on the rockburst prediction first analyzed the distribution features of microseismic data before a rockburst occurred. Then, the rockburst prediction indexes are presented, and the machine learning algorithm is used to describe the relationship between the rockburst prediction indexes and the future rockburst risks and types. However, these studies ignore the influence of the main controlling factors of rockbursts on the distribution of MS events before a rockburst occurs. In the rockburst prediction and classification process, there is no quantitative analysis of the influence of the intra-class variance and inter-class gap of prediction vectors on the prediction accuracy. The prediction vectors are the result of rockburst prediction indexes or MS data processing by deep learning encoders.</p>
<p>A quantitative information model of geology and mining is constructed based on the different types of rockburst risk events. The temporal and spatial distribution characteristics of MS data before a rockburst, in the main controlling factors of rockburst (i.e., geological structure, roof, coal pillar, and high-stress coal mass), are studied. The results show that the distribution features of microseismic and rockburst events of the same type may differ, and different types of events may show similarities in space-time distribution. The relationship between the performance of a deep learning prediction algorithm and the prediction vector of rockbursts is studied. The rockburst prediction model must accurately distinguish the differences between different types of data and accurately describe the intra-class variance of monitoring data.</p>
<p>A rockburst risk and type prediction method based on a convolutional neural network (CNN)-gated recurrent unit (GRU) and prototype learning prediction head is proposed. The CNN-GRU model can generate prediction vectors by integrating both implicit and explicit information derived from original microseismic (MS) data and prediction indexes. Through the use of cross-entropy loss, vector-prototype contrastive loss, and inter-vector-prototype contrastive loss, the model can autonomously manage the intra-class variance and inter-class distance of prediction vectors corresponding to different rockburst risks and types. Extensive experimental results demonstrate that our proposed CNN-GRU model with prototype-based prediction outperforms alternative algorithms in forecasting rockburst risks and types using MS data.</p>
</sec>
<sec id="s2">
<title>2 Spatial&#x2013;temporal and intensity evolution of rockbursts in MS data</title>
<p>Although the factors influencing rockbursts are complex, four controlling factors affecting the structure and stress of rockbursts are commonly recognized: geological structure, hard rock strata, coal pillars, and mining depth. The geological and mining data related to these four factors from typical rockburst events are first quantified. For each working face, the mining progress over a 10-day sliding window throughout the entire mining process is calculated. The geological and mining information models corresponding to different working faces are constructed. Then, by projecting the MS data to the quantified geological and mining information model, MS events and rockbursts are simply divided into four types, that is, the hard rock roof type, geological structure type, coal pillar type, and mining depth type, as shown in <xref ref-type="fig" rid="F1">Figure 1</xref>. Some events are caused by more than one controlling factor.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>The architecture of the CNN-GRU rockburst risk and type prediction model with a prototype-based prediction head.</p>
</caption>
<graphic xlink:href="feart-13-1601090-g001.tif"/>
</fig>
<p>
<xref ref-type="fig" rid="F2">Figure 2</xref> shows that the distribution of different types of MS events is disordered and irregular within 10 days before the rockburst. Even for the same type of events, microseismic events are clustered in several central regions and do not show similar distribution characteristics. Meanwhile, the same type of MS events is clustered in several central regions and do not show similar distribution characteristics. During long-term mining, most MS events show the characteristics of irregularity and relatively uniform distribution and occur in the area where the main control factors are active. When MS events occur intensively in some special areas, it indicates that there is a rockburst risk in the area. When MS events occur intensively in some special areas during the short term, this indicates that there is a rockburst risk in these areas.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Rockbursts and MS events caused by different main controlling factors within the working face. The MS data of all figures are within 10 days of the rockburst occurrence. The first-row figures show the dangerous MS events influenced by the <bold>(a)</bold> hard rock roof, <bold>(b)</bold> geological structure, <bold>(c)</bold> coal pillar, or <bold>(d)</bold> gravity or mining depth factors. The second row figures show the projection results of all MS data.</p>
</caption>
<graphic xlink:href="feart-13-1601090-g002.tif"/>
</fig>
<p>The sequential expansion of the time and MS energy across the entire working face is illustrated in <xref ref-type="fig" rid="F3">Figure 3</xref>. Over the long term, it is obvious that the main active controlling factors of a rockburst are the same as those of most MS events. This is because there are structural planes that hinder energy propagation among these four main controlling factors, and the occurrence area of most MS events can be used to determine the destination or main path of energy propagation. The variation of MS energy over time scales is disordered. From a short-term perspective within a rockburst occurrence, the energy and frequency of MS events have increased slightly, showing a relatively active state.</p>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Time sequence exhibition of various MS events within the working faces where a <bold>(a)</bold> hard rock roof type rockburst, <bold>(b)</bold> geological structure type rockburst, <bold>(c)</bold> coal pillar type rockburst, and <bold>(d)</bold> mining depth or gravity type rockburst occurs. Selected working faces are the same as those in <xref ref-type="fig" rid="F3">Figure 3</xref>. The MS event with the maximum energy value is the rockburst.</p>
</caption>
<graphic xlink:href="feart-13-1601090-g003.tif"/>
</fig>
<p>In summary, the MS events can be divided into four types according to the time and space projection of MS events on the controlling factors. For the entire mining process, all types of MS events are disordered and irregular in terms of temporal and spatial characteristics. When the same type of MS events evolves from disordered and scattered to ordered and intensive, the main controlling factor area may have a rockburst hazard. However, even in the short term of rockburst occurrence, MS events are clustered within several centers on the time and space scale. Therefore, it is difficult and complex to summarize the temporal and spatial characteristics of the same type of MS data and accurately distinguish between different types.</p>
</sec>
<sec id="s3">
<title>3 Relationship of prediction head and rockburst prediction vectors</title>
<p>In this paper, the content of rockburst prediction includes the risk and type. There are two rockburst risk prediction results: dangerous and non-dangerous. There are four rockburst type prediction results: gravity type, coal pillar type, roof type, and tectonic type. For the machine learning-based rockburst risk and type prediction method, the algorithms mainly consist of the encoder, decoder, and prediction head. The encoder is mainly responsible for analyzing the input MS data and extracting data features. For most previous rockburst prediction methods, the function of an encoder was replaced by the calculation formulas of prediction indexes. The decoder analyzes the feature space of the data features and processes them into vectors with rockburst risk and type information. The final prediction head maps such vectors to the probabilities of rockburst risk and rockburst types.</p>
<p>This section takes rockburst risk prediction as an example to study the distribution requirement of decoder output vectors and representative vectors in the prediction head. For the traditional rockburst risk prediction head, it is reasonable to assume that there are two representative vectors, that is, a non-dangerous representative vector w1 and a dangerous representative vector w2. After the rockburst prediction vector <italic>v</italic> is obtained, the essence of this problem is to determine the probability that <italic>v</italic> belongs to risk class <italic>C</italic> according to the relationship between <italic>v</italic> and the representative vector <italic>w</italic>
<sub>
<italic>c</italic>
</sub> in the prediction head, which can be regarded as follows:<disp-formula id="e1">
<mml:math id="m1">
<mml:mrow>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msubsup>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:mi>v</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>b</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>c</mml:mi>
<mml:mi>o</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>s</mml:mi>
<mml:mi>t</mml:mi>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>
</p>
<p>The least squares error between the dangerous probability calculated based on the vector <italic>v</italic> and the target yi can be written as follows:<disp-formula id="e2">
<mml:math id="m2">
<mml:mrow>
<mml:mi>L</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:mfrac>
<mml:mstyle displaystyle="true">
<mml:munderover>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
</mml:mstyle>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msubsup>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>b</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
<label>(2)</label>
</disp-formula>where N is the total number of samples. To simplify the calculation, the C label is changed to <italic>y</italic>
<sub>
<italic>c</italic>
</sub> &#x3d; <italic>N</italic>/<italic>N</italic>
<sub>1</sub>, and the other label is changed to <inline-formula id="inf1">
<mml:math id="m3">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>N</mml:mi>
<mml:mo>/</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>. To obtain the maximum likelihood estimate for <italic>b</italic>
<sub>
<italic>c</italic>
</sub>, setting <italic>L</italic> with respect to <italic>b</italic>
<sub>
<italic>c</italic>
</sub> to 0, the following result can be obtained:<disp-formula id="e3">
<mml:math id="m4">
<mml:mrow>
<mml:msub>
<mml:mi>b</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:msubsup>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(3)</label>
</disp-formula>where <italic>m</italic>
<sub>
<italic>c</italic>
</sub> and <inline-formula id="inf2">
<mml:math id="m5">
<mml:mrow>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> are the two category centers of dangerous and non-dangerous vectors. Setting the derivatives of <italic>L</italic> with respect to <italic>w</italic>
<sub>
<italic>c</italic>
</sub>:<disp-formula id="e4">
<mml:math id="m6">
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mi mathvariant="normal">&#x2202;</mml:mi>
<mml:mi>L</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi mathvariant="normal">&#x2202;</mml:mi>
<mml:msub>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:msubsup>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>T</mml:mi>
</mml:msubsup>
<mml:mo>&#x2b;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mover accent="true">
<mml:mi>C</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>m</mml:mi>
</mml:msub>
<mml:msubsup>
<mml:mi>v</mml:mi>
<mml:mi>m</mml:mi>
<mml:mi>T</mml:mi>
</mml:msubsup>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msup>
<mml:mi>m</mml:mi>
<mml:mi>T</mml:mi>
</mml:msup>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:msup>
<mml:mi>m</mml:mi>
<mml:mi>T</mml:mi>
</mml:msup>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:msub>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>N</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(4)</label>
</disp-formula>where <inline-formula id="inf3">
<mml:math id="m7">
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>/</mml:mo>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula>. In order to analyze the data distribution requirements of vector <italic>v</italic> data, the intra-class variance is<disp-formula id="e5">
<mml:math id="m8">
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2282;</mml:mo>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:msubsup>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x2b;</mml:mo>
</mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2282;</mml:mo>
<mml:mover accent="true">
<mml:mi>C</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:msubsup>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msub>
<mml:mi>S</mml:mi>
<mml:mi>w</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(5)</label>
</disp-formula>where <italic>k</italic> is the class index. The <italic>S</italic>
<sub>
<italic>W</italic>
</sub> is<disp-formula id="e6">
<mml:math id="m9">
<mml:mrow>
<mml:msub>
<mml:mi>S</mml:mi>
<mml:mi>W</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2282;</mml:mo>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mrow>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
<mml:mo>&#x2b;</mml:mo>
</mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2282;</mml:mo>
<mml:mover accent="true">
<mml:mi>C</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:mrow>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(6)</label>
</disp-formula>
</p>
<p>The inter-class distance can be expressed as<disp-formula id="e7">
<mml:math id="m10">
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msub>
<mml:mi>S</mml:mi>
<mml:mi>B</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
<label>(7)</label>
</disp-formula>where the <italic>S</italic>
<sub>
<italic>B</italic>
</sub> is<disp-formula id="e8">
<mml:math id="m11">
<mml:mrow>
<mml:msub>
<mml:mi>S</mml:mi>
<mml:mi>B</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:math>
<label>(8)</label>
</disp-formula>
</p>
<p>According to <xref ref-type="disp-formula" rid="e6">Equation 6</xref>, the following equation can be obtained.<disp-formula id="e9">
<mml:math id="m12">
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mover accent="true">
<mml:mi>C</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:msubsup>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:mo>&#x2b;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>m</mml:mi>
</mml:msub>
<mml:msubsup>
<mml:mi>v</mml:mi>
<mml:mi>m</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi>S</mml:mi>
<mml:mi>W</mml:mi>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msup>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:msup>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:math>
<label>(9)</label>
</disp-formula>
</p>
<p>Combining <xref ref-type="disp-formula" rid="e8">Equation 8</xref> and <xref ref-type="disp-formula" rid="e9">Equation 9</xref>, <xref ref-type="disp-formula" rid="e4">Equation 4</xref> can be transformed to<disp-formula id="e10">
<mml:math id="m13">
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mover accent="true">
<mml:mi>C</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:msubsup>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:mo>&#x2b;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mo>&#x2208;</mml:mo>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:munder>
</mml:mstyle>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>m</mml:mi>
</mml:msub>
<mml:msubsup>
<mml:mi>v</mml:mi>
<mml:mi>m</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi>S</mml:mi>
<mml:mi>W</mml:mi>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msup>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:msup>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:math>
<label>(10)</label>
</disp-formula>
</p>
<p>Setting <xref ref-type="disp-formula" rid="e10">Equation 10</xref> to zero, the following equation can be obtained:<disp-formula id="e11">
<mml:math id="m14">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3c9;</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x221d;</mml:mo>
<mml:msubsup>
<mml:mi>S</mml:mi>
<mml:mi>W</mml:mi>
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msubsup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(11)</label>
</disp-formula>
</p>
<p>Due to the differences in the monitoring data, SW is positive and not proportional to the unit matrix. Therefore, in the task of rockburst risk prediction, the vector <italic>v</italic> should not only maximize the variance between categories (<inline-formula id="inf4">
<mml:math id="m15">
<mml:mrow>
<mml:msub>
<mml:mi>m</mml:mi>
<mml:mover accent="true">
<mml:mi>c</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> and <italic>m</italic>
<sub>
<italic>c</italic>
</sub>) but also minimize the variance within each category SW. These two objects are interactive. This conclusion is also applicable to the prediction of rockburst type. If there is no proper training method, reducing the intra-class variance of vectors within the same class will also lead to the reduction of inter-class distance of vectors. Therefore, in this paper, the decoder output vectors of rockburst risk and type are controlled using a prototype-based prediction head that can accurately control the intra- and inter-class variances of vectors.</p>
</sec>
<sec id="s4">
<title>4 CNN-GRU model with the prototype-based prediction head</title>
<p>According to the analysis in <xref ref-type="sec" rid="s2">Section 2</xref> and <xref ref-type="sec" rid="s3">Section 3</xref>, the distribution of rockburst data is very complicated, and it is difficult to describe the distribution of rockburst prediction vectors belonging to one rockburst risk class or type by a representative vector in the traditional classifier head. Moreover, the traditional prediction headers based on single or multi-layer perceptrons are updated by backpropagation. This results in the distribution space of the rockburst prediction vectors that cannot be reasonably controlled. Therefore, a novel CNN-GRU algorithm for comprehensively analyzing the original MS data and rockburst prediction indexes is presented to project the MS data to rockburst risk and type prediction vectors by automatically analyzing and fusing the spatiotemporal distribution features of MS data. Then, a prototype-based rockburst risk and type prediction head that can control the inter-class distance and intra-class variance of the prediction vectors is constructed.</p>
<sec id="s4-1">
<title>4.1 The encoder-decoder based CNN-GRU algorithm</title>
<p>As mentioned in <xref ref-type="sec" rid="s2">Section 2</xref>, the distribution characteristics of MS events in different stages of rockburst development are very complicated. Therefore, in previous studies, the MS data are first transformed into prediction indexes before inputting the machine learning model. In this paper, the CNN-GRU algorithm outputs the prediction vectors by comprehensively analyzing the original MS data and the prediction indexes, as shown in <xref ref-type="fig" rid="F4">Figure 4</xref>.</p>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>The illustration of prototype updating and matching in the prediction head.</p>
</caption>
<graphic xlink:href="feart-13-1601090-g004.tif"/>
</fig>
<p>The encoder consists of five one-dimensional (1D) CNN-based layers. Each 1D CNN layer contains a convolution layer, a batch normalization layer, and a rectified linear unit (ReLU) activation function. The number of convolution kernels in each layer is 12, 12, 12, 24, and 24, respectively. The input of the CNN-based encoder is the original MS data with a standardized time interval of two adjacent events, X, Y, and Z coordinates, and energy. If the number of analyzed MS data elements is n and the dimension of the prediction vector is 24, the input and output dimensions of the encoder are [batch size, n, 5] and [1, batch size, 24].</p>
<p>The decoder consists of five GRU-based modules. The hidden layer parameter of the first GRU module is the encoder output. The input of each GRU module is made up of different rockburst prediction indexes. The decoder outputs are the prediction vectors whose dimension is [1, batch size, 24]. In this model, the rockburst prediction index consists of standardized temporal concentration <italic>Q</italic>
<sub>
<italic>T</italic>
</sub>, the time information entropy <italic>Q</italic>
<sub>
<italic>t</italic>
</sub>, space concentration <italic>Q</italic>
<sub>
<italic>D</italic>
</sub>, spatiotemporal diffusion <italic>d</italic>
<sub>
<italic>s</italic>
</sub>, and energy concentration index <italic>Q</italic>
<sub>
<italic>E</italic>
</sub>. The temporal concentration <italic>Q</italic>
<sub>
<italic>T</italic>
</sub> can be formalized as<disp-formula id="e12">
<mml:math id="m16">
<mml:mrow>
<mml:msub>
<mml:mi>Q</mml:mi>
<mml:mi>T</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>V</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>r</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>T</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>/</mml:mo>
<mml:mo>&#x394;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>T</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
<label>(12)</label>
</disp-formula>where <italic>Q</italic>
<sub>
<italic>T</italic>
</sub>,<italic>Var</italic>(<italic>T</italic>
<sub>
<italic>n</italic>
</sub>) and <inline-formula id="inf5">
<mml:math id="m17">
<mml:mrow>
<mml:mo>&#x394;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>T</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> are the temporal concentration, the variance, and the mean value of the time interval of the last <italic>n</italic> MS events.</p>
<p>Therefore, the time information entropy <italic>Q</italic>
<sub>
<italic>t</italic>
</sub> is used to describe the aggregation degree of MS events in the time series, reflecting the disorder or order in the evolution of MS time.<disp-formula id="e13">
<mml:math id="m18">
<mml:mrow>
<mml:msub>
<mml:mi>Q</mml:mi>
<mml:mi>t</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>/</mml:mo>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:msubsup>
</mml:mstyle>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2061;</mml:mo>
<mml:msup>
<mml:mi>ln</mml:mi>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:msup>
</mml:mrow>
<mml:msup>
<mml:mi>ln</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:msup>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(13)</label>
</disp-formula>where <italic>n</italic> is the total number of selected MS events. <inline-formula id="inf6">
<mml:math id="m19">
<mml:mrow>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:msub>
<mml:mi>t</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>t</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi>t</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>t</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
</inline-formula>, <italic>t</italic>
<sub>
<italic>i</italic>
</sub> is the occurrence time of the i-th MS event, and the value of <italic>p</italic>
<sub>
<italic>i</italic>
</sub> is 0&#x223c;1.</p>
<p>The spatial distribution of MS events is important for understanding the stability of the coal rock mass in the mine.<disp-formula id="e14">
<mml:math id="m20">
<mml:mrow>
<mml:msub>
<mml:mi>Q</mml:mi>
<mml:mi>D</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>V</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>r</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>R</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>/</mml:mo>
<mml:mo>&#x394;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>R</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
<label>(14)</label>
</disp-formula>where <italic>Q</italic>
<sub>
<italic>D</italic>
</sub>, <italic>Var</italic>(<italic>R</italic>
<sub>
<italic>n</italic>
</sub>), and <inline-formula id="inf7">
<mml:math id="m21">
<mml:mrow>
<mml:mo>&#x394;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>R</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> are the space concentration, variance, and mean value of the radius corresponding to the last <italic>n</italic> MS events.</p>
<p>The spatiotemporal diffusion is summarized to reflect the dispersion degree of MS events in time and space.<disp-formula id="e15">
<mml:math id="m22">
<mml:mrow>
<mml:msub>
<mml:mi>d</mml:mi>
<mml:mi>s</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mover accent="true">
<mml:mi>X</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>/</mml:mo>
<mml:mover accent="true">
<mml:mi>t</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(15)</label>
</disp-formula>where <inline-formula id="inf8">
<mml:math id="m23">
<mml:mrow>
<mml:mover accent="true">
<mml:mi>X</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula> is the average distance between sequential MS events. <inline-formula id="inf9">
<mml:math id="m24">
<mml:mrow>
<mml:mover accent="true">
<mml:mi>t</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula> is the average time interval between sequential MS events.</p>
<p>The energy concentration index is established to reflect the energy change and MS distribution before the rockburst.<disp-formula id="e16">
<mml:math id="m25">
<mml:mrow>
<mml:msub>
<mml:mi>Q</mml:mi>
<mml:mi>E</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>V</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>r</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>E</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>/</mml:mo>
<mml:mo>&#x394;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>E</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
<label>(16)</label>
</disp-formula>where <italic>Q</italic>
<sub>
<italic>E</italic>
</sub>, <italic>Var</italic>(<italic>E</italic>
<sub>
<italic>n</italic>
</sub>), and <inline-formula id="inf10">
<mml:math id="m26">
<mml:mrow>
<mml:mo>&#x394;</mml:mo>
<mml:msub>
<mml:mover accent="true">
<mml:mi>E</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> are the energy concentration, the variance, and the mean value of the energy corresponding to the continuous n MS events.</p>
</sec>
<sec id="s4-2">
<title>4.2 The prototype-based rockburst risk and type prediction head</title>
<p>To control the inter-class distance and intra-class variance of the prediction vectors, a novel rockburst risk and type prediction head that can constantly adjust the position of the prototype is proposed, as seen in <xref ref-type="fig" rid="F4">Figure 4</xref>. In the perceptron-based prediction head, the prototype can be regarded as the representative vector in the traditional prediction head.</p>
<p>In this method, the sub-centers of vectors belonging to class <italic>C</italic> are described by <italic>K</italic> prototypes, that is, <inline-formula id="inf11">
<mml:math id="m27">
<mml:mrow>
<mml:msubsup>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>K</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula>. <italic>P</italic>
<sub>
<italic>c</italic>,<italic>k</italic>
</sub> is the <italic>k</italic>-th sub-cluster center of prediction vectors belonging to class <italic>C</italic>. With this prototype-based classifier, the probability distribution of prediction vector <italic>v</italic> over the <italic>C</italic> class can be described as<disp-formula id="e17">
<mml:math id="m28">
<mml:mrow>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>v</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>v</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>c</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:msubsup>
<mml:mi mathvariant="normal">&#x3a3;</mml:mi>
<mml:mrow>
<mml:msup>
<mml:mi>c</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>c</mml:mi>
</mml:msubsup>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>v</mml:mi>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mi>c</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
<mml:mtext>with&#x2009;</mml:mtext>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>v</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>c</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>min</mml:mi>
<mml:mo>&#x2061;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mfenced open="&#x2329;" close="&#x232a;" separators="|">
<mml:mrow>
<mml:mi>v</mml:mi>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>K</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:math>
<label>(17)</label>
</disp-formula>where the vector-class distance <italic>s</italic>
<sub>
<italic>v,c</italic>
</sub> &#x2208; [&#x2212;1, 1] is the distance to the closest prototype of class <italic>c</italic>. Based on <xref ref-type="disp-formula" rid="e12">Equation 12</xref>, the cross-entropy loss is<disp-formula id="e18">
<mml:math id="m29">
<mml:mrow>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:mi>C</mml:mi>
<mml:mi>E</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>log</mml:mi>
<mml:mfrac>
<mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>v</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>c</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:msubsup>
<mml:mi mathvariant="normal">&#x3a3;</mml:mi>
<mml:mrow>
<mml:msup>
<mml:mi>c</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>c</mml:mi>
</mml:msubsup>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>s</mml:mi>
<mml:mrow>
<mml:mi>v</mml:mi>
<mml:mo>,</mml:mo>
<mml:msup>
<mml:mi>c</mml:mi>
<mml:mo>&#x2032;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(18)</label>
</disp-formula>
</p>
<p>As studied in <xref ref-type="sec" rid="s3">Section 3</xref>, the prediction head should push the prediction vector close to a certain prototype (i.e., the center of vectors) of class <italic>c</italic> and distant from other prototypes belonging to other classes. However, <xref ref-type="disp-formula" rid="e7">Equation 7</xref> only considers one vector-class distance. Therefore, the updating method and other limiting conditions of the prototype should be studied.</p>
<p>The prototypes are selected and assigned by the online clustering. Prediction vectors within the same class are assigned to the prototypes of that class, and the prototypes are then updated according to the assignments. Formally, given the vectors <inline-formula id="inf12">
<mml:math id="m30">
<mml:mrow>
<mml:msub>
<mml:mi>V</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> in a training batch that belongs to class <italic>c</italic>, the goal is to map the vectors <italic>V</italic>
<sub>C</sub> to the <italic>K</italic> prototypes <inline-formula id="inf13">
<mml:math id="m31">
<mml:mrow>
<mml:msub>
<mml:mi>P</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>K</mml:mi>
</mml:msubsup>
</mml:mrow>
</mml:math>
</inline-formula> of class <italic>C</italic>. The vector-to-prototype mapping is denoted as <inline-formula id="inf14">
<mml:math id="m32">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:msubsup>
<mml:mo>&#x2208;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>K</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>. <inline-formula id="inf15">
<mml:math id="m33">
<mml:mrow>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:msubsup>
<mml:mrow>
<mml:mfenced open="[" close="]" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>k</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>K</mml:mi>
</mml:msubsup>
<mml:mo>&#x2208;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi>K</mml:mi>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> is the one-hot assignments of vectors <italic>v</italic>
<sub>
<italic>n</italic>
</sub> over the <italic>K</italic> prototypes. The optimization of <italic>L</italic>
<sub>C</sub> is achieved by maximizing the similarity between vector embeddings and the prototypes.<disp-formula id="e19">
<mml:math id="m34">
<mml:mrow>
<mml:mtable columnalign="left">
<mml:mtr>
<mml:mtd>
<mml:mrow>
<mml:munder>
<mml:mi>max</mml:mi>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:munder>
<mml:mtext>&#x2009;Tr</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msubsup>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msubsup>
<mml:mi>P</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msub>
<mml:mi>V</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:mtd>
</mml:mtr>
<mml:mtr>
<mml:mtd>
<mml:mrow>
<mml:mi>s</mml:mi>
<mml:mo>.</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>.</mml:mo>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mfenced open="{" close="}" separators="|">
<mml:mrow>
<mml:mn>0</mml:mn>
<mml:mo>,</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mrow>
<mml:mi>K</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msup>
<mml:mo>,</mml:mo>
<mml:msubsup>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>K</mml:mi>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:msup>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>K</mml:mi>
</mml:msup>
</mml:mrow>
</mml:mtd>
</mml:mtr>
</mml:mtable>
</mml:mrow>
</mml:math>
<label>(19)</label>
</disp-formula>
</p>
<p>The unique assignment constraint <inline-formula id="inf16">
<mml:math id="m35">
<mml:mrow>
<mml:msubsup>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>K</mml:mi>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> ensures that each vector is assigned to one and only one prototype. The equipartition constraint <inline-formula id="inf17">
<mml:math id="m36">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>K</mml:mi>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> enforces that each prototype is selected at least N/K times in the batch on average. This constraint prevents all vectors from being assigned to a single prototype and eventually benefits the representative ability of the prototypes. To solve <xref ref-type="disp-formula" rid="e8">Equation 8</xref>, Lc can be relaxed to an element of the transportation polytope:<disp-formula id="e20">
<mml:math id="m37">
<mml:mrow>
<mml:mtable columnalign="left">
<mml:mtr>
<mml:mtd>
<mml:mrow>
<mml:munder>
<mml:mi>max</mml:mi>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:munder>
<mml:mtext>&#x2009;Tr</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msubsup>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msubsup>
<mml:mi>P</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msub>
<mml:mi>V</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:mi>&#x3ba;</mml:mi>
<mml:mi>h</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:mtd>
</mml:mtr>
<mml:mtr>
<mml:mtd>
<mml:mrow>
<mml:mi>s</mml:mi>
<mml:mo>.</mml:mo>
<mml:mi>t</mml:mi>
<mml:mo>.</mml:mo>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x2208;</mml:mo>
<mml:msubsup>
<mml:mi mathvariant="normal">R</mml:mi>
<mml:mo>&#x2b;</mml:mo>
<mml:mrow>
<mml:mi>K</mml:mi>
<mml:mo>&#xd7;</mml:mo>
<mml:mi>N</mml:mi>
</mml:mrow>
</mml:msubsup>
<mml:mo>,</mml:mo>
<mml:msubsup>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>K</mml:mi>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:msup>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>N</mml:mi>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>N</mml:mi>
</mml:mrow>
<mml:mrow>
<mml:mi>K</mml:mi>
</mml:mrow>
</mml:mfrac>
<mml:msup>
<mml:mn>1</mml:mn>
<mml:mi>K</mml:mi>
</mml:msup>
</mml:mrow>
</mml:mtd>
</mml:mtr>
</mml:mtable>
</mml:mrow>
</mml:math>
<label>(20)</label>
</disp-formula>where <inline-formula id="inf18">
<mml:math id="m38">
<mml:mrow>
<mml:mi>h</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>log</mml:mi>
<mml:mo>&#x2061;</mml:mo>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>v</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is an entropy, and <italic>&#x3ba;</italic> &#x3e; 0 is a parameter that controls the smoothness of the distribution. With the soft assignment relaxation and the extra regularization term h (Lc), the solution of <xref ref-type="disp-formula" rid="e9">Equation 9</xref> can be given as<disp-formula id="e21">
<mml:math id="m39">
<mml:mrow>
<mml:msub>
<mml:mi>L</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mtext>diag</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>&#x3c8;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:msubsup>
<mml:mi>P</mml:mi>
<mml:mi>c</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msubsup>
<mml:msub>
<mml:mi>V</mml:mi>
<mml:mi>c</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mi>&#x3ba;</mml:mi>
</mml:mfrac>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mtext>diag</mml:mtext>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mi>&#x3b5;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:math>
<label>(21)</label>
</disp-formula>where <italic>&#x3c8;</italic>&#x2208;R<sup>
<italic>K</italic>
</sup> and <italic>&#x3b5;</italic>&#x2208;R<sup>
<italic>N</italic>
</sup> are re-normalization vectors, computed by a few steps of the Sinkhorn&#x2013;Knopp iteration.</p>
<p>The vector-prototype contrastive learning strategy is employed to maximize the prototype assignment posterior probability and the variation of different prototypes. The vector-prototype contrastive loss is<disp-formula id="e22">
<mml:math id="m40">
<mml:mrow>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mi>P</mml:mi>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>log</mml:mi>
<mml:mfrac>
<mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msup>
<mml:mi>v</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>c</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>k</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
</mml:mrow>
</mml:msub>
<mml:mo>/</mml:mo>
<mml:mi>&#x3c4;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msup>
<mml:mi>v</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>c</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>k</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
</mml:mrow>
</mml:msub>
<mml:mo>/</mml:mo>
<mml:mi>&#x3c4;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:msup>
<mml:mi>p</mml:mi>
<mml:mo>&#x2212;</mml:mo>
</mml:msup>
<mml:mo>&#x2208;</mml:mo>
<mml:msup>
<mml:mi>P</mml:mi>
<mml:mo>&#x2212;</mml:mo>
</mml:msup>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2061;</mml:mo>
<mml:mi>exp</mml:mi>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:msup>
<mml:mi>v</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
<mml:msup>
<mml:mi>p</mml:mi>
<mml:mo>&#x2212;</mml:mo>
</mml:msup>
<mml:mo>/</mml:mo>
<mml:mi>&#x3c4;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:math>
<label>(22)</label>
</disp-formula>where the temperature <italic>&#x3c4;</italic> controls the concentration level of representations. <xref ref-type="disp-formula" rid="e11">Equation 11</xref> enforces each vector <italic>v</italic> to be similar with its assigned prototype <inline-formula id="inf19">
<mml:math id="m41">
<mml:mrow>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>c</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>k</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula>, and dissimilar with other irrelevant prototypes <italic>p</italic>
<sup>&#x2212;</sup>.</p>
<p>
<xref ref-type="disp-formula" rid="e18">Equation 18</xref> and <xref ref-type="disp-formula" rid="e22">Equation 22</xref> inspire inter-class and inter-cluster discrimination but do not consider reducing the intra-cluster variation, that is, making vectors of the same prototype compact. Thus, a compactness-aware loss is employed for further regularizing representations by directly minimizing the distance between each prediction vector and its assigned prototype:<disp-formula id="e23">
<mml:math id="m42">
<mml:mrow>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mi>P</mml:mi>
<mml:mi>D</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:msup>
<mml:mi>v</mml:mi>
<mml:mo>&#x22a4;</mml:mo>
</mml:msup>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>c</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>k</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
<label>(23)</label>
</disp-formula>
</p>
<p>Note that both <italic>v</italic> and <inline-formula id="inf20">
<mml:math id="m43">
<mml:mrow>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:msub>
<mml:mi>c</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
<mml:mo>,</mml:mo>
<mml:msub>
<mml:mi>k</mml:mi>
<mml:mi>v</mml:mi>
</mml:msub>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> are 2-normalized. To accurately control the results, the prototypes were not learned by stochastic gradient descent. After each training iteration, each prototype is updated as:<disp-formula id="e24">
<mml:math id="m44">
<mml:mrow>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2190;</mml:mo>
<mml:mi>&#x3bc;</mml:mi>
<mml:msub>
<mml:mi>p</mml:mi>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mrow>
<mml:mfenced open="(" close=")" separators="|">
<mml:mrow>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mi>&#x3bc;</mml:mi>
</mml:mrow>
</mml:mfenced>
</mml:mrow>
<mml:msub>
<mml:mover accent="true">
<mml:mi>v</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
<label>(24)</label>
</disp-formula>where the updated weight <italic>&#x3bc;</italic> is 0.999. <inline-formula id="inf21">
<mml:math id="m45">
<mml:mrow>
<mml:msub>
<mml:mover accent="true">
<mml:mi>v</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
<mml:mrow>
<mml:mi>c</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>k</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the mean vector of the embedded training pixels, which are assigned to prototype <italic>p</italic>
<sub>
<italic>c, k</italic>
</sub> by online clustering.</p>
<p>This training objective minimizes intra-cluster variations while maintaining separation between features with different prototype assignments. As shown in <xref ref-type="fig" rid="F5">Figure 5</xref>, the total loss for prototype-based prediction head is<disp-formula id="e25">
<mml:math id="m46">
<mml:mrow>
<mml:mi>L</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi>&#x3bb;</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:mi>C</mml:mi>
<mml:mi>E</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>&#x3bb;</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mi>P</mml:mi>
<mml:mi>C</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>&#x3bb;</mml:mi>
<mml:mn>3</mml:mn>
</mml:msub>
<mml:msub>
<mml:mi>l</mml:mi>
<mml:mrow>
<mml:mi>P</mml:mi>
<mml:mi>P</mml:mi>
<mml:mi>D</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
<label>(25)</label>
</disp-formula>where <italic>&#x3bb;</italic>
<sub>1</sub>&#x223c;<italic>&#x3bb;</italic>
<sub>3</sub> are the weights for the loss function. Their values are 1, 0.5, and 0.1 in the article.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>The illustration of different loss functions for the prototype-based prediction head.</p>
</caption>
<graphic xlink:href="feart-13-1601090-g005.tif"/>
</fig>
<p>As shown in <xref ref-type="fig" rid="F1">Figure 1</xref>, there are two outputs of the prototype-based prediction head, that is, the rockburst types and rockburst risks. There are four classes for the rockburst types, including geological structure, hard rock strata, coal pillars, and mining depth. The class of rockburst risks consists of dangerous and non-dangerous events.</p>
</sec>
</sec>
<sec id="s5">
<title>5 Engineering application and experiment implementation details</title>
<sec id="s5-1">
<title>5.1 Rockburst type and risk prediction data</title>
<p>In the experiments, microseismic data from three different rockburst-prone mines are utilized, with the number of monitoring data elements being 29,156, 48,318, and 32,796, respectively. The geological formations and mining conditions exhibit significant heterogeneity for these three mines. Because these mines are located in different places, the geological and mining conditions vary among these rockburst mines. The MS data information includes time interval, X, Y, and Z coordinates, and energy value, respectively. The types of rockburst events and MS events are determined by the projection results of events in different rockburst controlling factor areas. The rockburst types are geological structure, hard rock strata, coal pillars, and mining depth. The rockburst risk levels are dangerous and non-dangerous. In the original MS data, the proportion of rockburst risk events is very low. Therefore, in the training and testing stages, large energy events, obvious mine earthquakes, and rockbursts are defined as dangerous events, and other MS events are defined as non-dangerous samples. In the presented prototype-based prediction head, the prototype numbers for each rockburst type and risk level are set as 5.</p>
<p>In this paper, the standardized time interval of two adjacent events, the X, Y, and Z coordinates, and the energy are selected as the input features of the encoder. The standardized rockburst prediction indexes are selected as the input feature of the decoder. The target is future MS event types and risk levels. By employing the inputs as features and future MS event types and rockburst risk levels as labels, the samples for rockburst types and levels prediction are constructed. When processing the original MS data to sample features, by setting the fused number of MS events as <italic>n</italic> &#x3d; 5, 6, &#x2026;, 15, thirty datasets for each coal mine are constructed. The partition ratio of training and test samples is 7:3 based on the continuous timeline. Considering the imbalance between different risk level events as well as the number of dangerous and non-dangerous samples, the weighted sampling method is used during training.</p>
</sec>
<sec id="s5-2">
<title>5.2 Comparison methods and implementation details</title>
<p>The comparison methods are SVM (<xref ref-type="bibr" rid="B15">Ji et al., 2020</xref>), CNN (<xref ref-type="bibr" rid="B34">Zhang et al., 2021</xref>), LSTM (<xref ref-type="bibr" rid="B10">Di et al., 2023b</xref>), and CNN-GRU with traditional prediction head methods, which are the most popular supervised machine learning methods for rockburst risk and type prediction. The traditional prediction head is a two-layer MLP network.</p>
<p>The training epoch for each method is 200. For the comparison method with the traditional prediction head, the CE loss function is employed to evaluate the difference between the model output and the targets. The input of comparison methods is the combination of sample features. All experiments are implemented on Pytorch 1.10.2 &#x2b; CUDA 11.3, in FP32 precision by using two RTX A6000 GPUs. A batch size of 512, an Adam optimizer with a momentum of 0.9, a weight decay of 4 &#xd7; 10<sup>&#x2212;5</sup>, and an initial learning rate of 5 &#xd7; 10<sup>&#x2212;4</sup> are employed.</p>
<p>To show the performance of different algorithms, the evaluation metric of different algorithms adopts the prediction accuracy of different rockburst types (Types), non-dangerous samples (Non), and dangerous samples (Dan) on the testing dataset.</p>
</sec>
<sec id="s5-3">
<title>5.3 Comparison results of different rockburst types and risk prediction methods</title>
<p>The comparison study results are shown in <xref ref-type="fig" rid="F6">Figure 6</xref>. The presented CNN-GRU model with the prototype-based prediction head shows the best performance. The accuracy of the proposed method on the testing set with different rockburst type samples, only non-dangerous samples, only dangerous events achieves 80.72%, 78.62%, and 82.57% in Mine 1, 78.19%, 89.13%, and 79.72% in Mine 2, and 76.07%, 86.92%, and 78.37% in Mine 3.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>The performance comparison of the presented method and other methods. The presented CNN-GRU method is the presented CNN-GRU model with the prototype-based prediction head. The traditional CNN-GRU method is the CNN-GRU model with a two-layer MLP prediction head.</p>
</caption>
<graphic xlink:href="feart-13-1601090-g006.tif"/>
</fig>
<p>When different models reach the highest accuracy on different data sets, the corresponding fused MS event number <italic>n</italic> is similar. For Mine 1 and Mine 2, when the fused MS event number <italic>n</italic> is 8, almost all methods achieve the best performance. For Mine 3, the best performance of the entire method appears when the fused MS event number <italic>n</italic> is 7. This may be because the rockburst prediction indexes and MS data are mainly physical quantity indexes, considering the physical logic of coal rock mass failure and mining rate. For the same mine or workface, the physical logic and mining rate are similar.</p>
</sec>
</sec>
<sec id="s6">
<title>6 Ablation experiments and discussion</title>
<sec id="s6-1">
<title>6.1 The number of prototypes</title>
<p>The number of prototypes in the prediction head influences the accurate description of the intra-class variance and inter-class distance of the rockburst prediction vectors. It can also affect the output distribution of the prediction vector through the loss function. The results of setting the fused MS event number <italic>n</italic> to 8/8/7 for Mine 1, Mine 2, and Mine 3 are shown in <xref ref-type="table" rid="T1">Table 1</xref>. The optimal number of prototypes per class is five, considering the trade-off between performance and computation cost. When the number of prototypes is too small, it is impossible to describe the variance between vectors in different classes. Too many prototypes can result in overfitting and reduced computational efficiency.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>The influence of prototype number on the presented CNN-GRU model with a prototype-based prediction head.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="2" align="center">Number</th>
<th colspan="3" align="center">Accuracy in Mine 1</th>
<th colspan="3" align="center">Accuracy in Mine 2</th>
<th colspan="3" align="center">Accuracy in Mine 3</th>
</tr>
<tr>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">1</td>
<td align="center">77.89</td>
<td align="center">87.69</td>
<td align="center">76.88</td>
<td align="center">75.23</td>
<td align="center">85.03</td>
<td align="center">74.36</td>
<td align="center">74.10</td>
<td align="center">82.77</td>
<td align="center">71.11</td>
</tr>
<tr>
<td align="center">3</td>
<td align="center">80.86</td>
<td align="center">90.06</td>
<td align="center">78.92</td>
<td align="center">78.22</td>
<td align="center">87.67</td>
<td align="center">76.60</td>
<td align="center">76.65</td>
<td align="center">85.02</td>
<td align="center">74.00</td>
</tr>
<tr>
<td align="center">5</td>
<td align="center">82.57</td>
<td align="center">91.62</td>
<td align="center">80.72</td>
<td align="center">79.72</td>
<td align="center">89.13</td>
<td align="center">78.19</td>
<td align="center">78.37</td>
<td align="center">86.92</td>
<td align="center">76.07</td>
</tr>
<tr>
<td align="center">7</td>
<td align="center">82.70</td>
<td align="center">91.82</td>
<td align="center">80.84</td>
<td align="center">79.88</td>
<td align="center">89.29</td>
<td align="center">78.38</td>
<td align="center">78.48</td>
<td align="center">87.08</td>
<td align="center">76.25</td>
</tr>
<tr>
<td align="center">10</td>
<td align="center">82.88</td>
<td align="center">91.93</td>
<td align="center">81.04</td>
<td align="center">80.07</td>
<td align="center">89.47</td>
<td align="center">78.48</td>
<td align="center">78.59</td>
<td align="center">87.26</td>
<td align="center">76.40</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s6-2">
<title>6.2 Different deep learning methods with the prototype-based prediction head</title>
<p>The prototype-based rockburst type and risk prediction head showed significant performance gain for the CNN-GRU method, as shown in <xref ref-type="fig" rid="F6">Figure 6</xref>. To demonstrate the universality and superiority of the prototype-based rockburst type and risk prediction head, the prototype-based prediction head is employed with other methods. The fused MS event number <italic>n</italic> is set as 8/8/7 for Mine 1, Mine 2, and Mine 3. <xref ref-type="table" rid="T2">Table 2</xref> reports that the prototype-based prediction head is obviously superior to the traditional prediction head for the employed deep learning methods. This is mainly because the traditional prediction head only uses the cross-entropy loss function to train, which cannot prevent the representative vectors of the same class from becoming similar in the training process. This results in the prediction head being unable to describe the inner-class variance even with a sufficient number of representative vectors. At the same time, it is impossible to control the encoder&#x2013;decoder to ensure the same kind of output prediction vectors close in the distribution space for the traditional prediction head. These problems are overcome in the prototype-based prediction head.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>The influence of prototype number on the presented CNN-GRU model with the prototype-based prediction head. CNN-GRU-Prototype and CNN-GRU-Tradition are the CNN-GRU model with a prototype-based prediction head and a traditional prediction head, respectively.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="2" align="center">Prediction head</th>
<th rowspan="2" align="center">Throughput (samples/s)</th>
<th colspan="3" align="center">Accuracy in Mine 1</th>
<th colspan="3" align="center">Accuracy in Mine 2</th>
<th colspan="3" align="center">Accuracy in Mine 3</th>
</tr>
<tr>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">CNN-Prototype</td>
<td align="center">420.8</td>
<td align="center">77.81</td>
<td align="center">86.80</td>
<td align="center">75.95</td>
<td align="center">74.27</td>
<td align="center">84.68</td>
<td align="center">73.07</td>
<td align="center">72.77</td>
<td align="center">81.02</td>
<td align="center">70.15</td>
</tr>
<tr>
<td align="center">CNN-Tradition</td>
<td align="center">429.7</td>
<td align="center">75.23</td>
<td align="center">83.99</td>
<td align="center">73.11</td>
<td align="center">72.07</td>
<td align="center">81.80</td>
<td align="center">70.83</td>
<td align="center">70.52</td>
<td align="center">78.94</td>
<td align="center">68.10</td>
</tr>
<tr>
<td align="center">LSTM-Prototype</td>
<td align="center">453.6</td>
<td align="center">79.87</td>
<td align="center">88.43</td>
<td align="center">77.40</td>
<td align="center">75.76</td>
<td align="center">85.94</td>
<td align="center">74.97</td>
<td align="center">74.21</td>
<td align="center">83.16</td>
<td align="center">72.34</td>
</tr>
<tr>
<td align="center">LSTM-Tradition</td>
<td align="center">462.1</td>
<td align="center">76.48</td>
<td align="center">85.18</td>
<td align="center">74.31</td>
<td align="center">73.05</td>
<td align="center">82.97</td>
<td align="center">71.98</td>
<td align="center">71.70</td>
<td align="center">80.44</td>
<td align="center">69.57</td>
</tr>
<tr>
<td align="center">CNN-GRU-Prototype</td>
<td align="center">405.9</td>
<td align="center">82.57</td>
<td align="center">91.62</td>
<td align="center">80.72</td>
<td align="center">79.72</td>
<td align="center">89.13</td>
<td align="center">78.19</td>
<td align="center">78.37</td>
<td align="center">86.92</td>
<td align="center">76.07</td>
</tr>
<tr>
<td align="center">CNN-GRU-Tradition</td>
<td align="center">393.3</td>
<td align="center">78.16</td>
<td align="center">87.19</td>
<td align="center">76.29</td>
<td align="center">74.96</td>
<td align="center">84.76</td>
<td align="center">73.78</td>
<td align="center">73.38</td>
<td align="center">82.44</td>
<td align="center">71.53</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s6-3">
<title>6.3 The weight of loss functions</title>
<p>There are three loss functions in <xref ref-type="disp-formula" rid="e25">Equation 25</xref>. Ablation studies conducted to evaluate their individual contributions, as shown in <xref ref-type="table" rid="T3">Table 3</xref>, demonstrate that all three components enhance model performance, albeit with varying impacts. The cross-entropy loss exhibits the most significant contribution, as it is the primary training objective. Meanwhile, the vector-prototype contrastive loss and the loss compactness-aware loss primarily regulate prototype distribution within the classifier. Notably, the experimental results reveal that maximizing the posterior probability of prototype assignments and enhancing prototype diversity between clusters are more critical than minimizing intra-cluster compactness.</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>The influence of loss function weights on the presented CNN-GRU model. The employed model is CNN-GRU-Prototype. <italic>&#x3bb;</italic>
<sub>1</sub>&#x223c;<italic>&#x3bb;</italic>
<sub>3</sub> are the loss function weights for cross-entropy loss, vector-prototype contrastive loss, and loss compactness-aware loss in <xref ref-type="disp-formula" rid="e25">Equation 25</xref>.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th colspan="3" align="center">Loss function weights</th>
<th colspan="3" align="center">Accuracy in Mine 1</th>
<th colspan="3" align="center">Accuracy in Mine 2</th>
<th colspan="3" align="center">Accuracy in Mine 3</th>
</tr>
<tr>
<th align="center">
<italic>&#x3bb;</italic>
<sub>1</sub>
</th>
<th align="center">
<italic>&#x3bb;</italic>
<sub>2</sub>
</th>
<th align="center">
<italic>&#x3bb;</italic>
<sub>3</sub>
</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">1</td>
<td align="center">0.5</td>
<td align="center">0.1</td>
<td align="center">82.57</td>
<td align="center">91.62</td>
<td align="center">80.72</td>
<td align="center">79.72</td>
<td align="center">89.13</td>
<td align="center">78.19</td>
<td align="center">78.37</td>
<td align="center">86.92</td>
<td align="center">76.07</td>
</tr>
<tr>
<td align="center">1</td>
<td align="center">0</td>
<td align="center">0.1</td>
<td align="center">80.13</td>
<td align="center">88.46</td>
<td align="center">78.12</td>
<td align="center">76.23</td>
<td align="center">86.79</td>
<td align="center">75.48</td>
<td align="center">76.89</td>
<td align="center">84.26</td>
<td align="center">73.82</td>
</tr>
<tr>
<td align="center">1</td>
<td align="center">0.5</td>
<td align="center">0</td>
<td align="center">81.27</td>
<td align="center">90.12</td>
<td align="center">78.94</td>
<td align="center">77.53</td>
<td align="center">87.05</td>
<td align="center">77.56</td>
<td align="center">77.24</td>
<td align="center">85.83</td>
<td align="center">75.06</td>
</tr>
<tr>
<td align="center">1</td>
<td align="center">0.5</td>
<td align="center">0.5</td>
<td align="center">81.92</td>
<td align="center">90.53</td>
<td align="center">79.64</td>
<td align="center">78.96</td>
<td align="center">88.06</td>
<td align="center">77.29</td>
<td align="center">77.10</td>
<td align="center">85.82</td>
<td align="center">74.95</td>
</tr>
<tr>
<td align="center">1</td>
<td align="center">1</td>
<td align="center">0.1</td>
<td align="center">82.01</td>
<td align="center">90.91</td>
<td align="center">80.15</td>
<td align="center">79.36</td>
<td align="center">88.23</td>
<td align="center">77.59</td>
<td align="center">77.63</td>
<td align="center">86.01</td>
<td align="center">75.26</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s6-4">
<title>6.4 The weight of loss functions</title>
<p>To evaluate the influence of input indicators on model performance, we conducted ablation experiments on five indicators input to the GRU model. The results are presented in <xref ref-type="table" rid="T4">Table 4</xref>. The experiments revealed that the energy concentration index contributed most significantly to model performance, while the time concentration index had the least impact. This phenomenon may be attributed to the fact that the energy concentration index directly reflects the accumulation and release processes of elastic strain energy within the rock mass. When energy becomes highly concentrated in a localized area, it indicates sufficient accumulation of strain energy, which may suddenly release upon reaching critical conditions, potentially triggering a rockburst.</p>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>The influence of input indexes on the presented CNN-GRU model.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th rowspan="2" align="center">The missing index</th>
<th colspan="3" align="center">Accuracy in Mine 1</th>
<th colspan="3" align="center">Accuracy in Mine 2</th>
<th colspan="3" align="center">Accuracy in Mine 3</th>
</tr>
<tr>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
<th align="center">Types</th>
<th align="center">Non</th>
<th align="center">Dan</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="center">Temporal concentration</td>
<td align="center">60.21</td>
<td align="center">69.21</td>
<td align="center">58.20</td>
<td align="center">57.23</td>
<td align="center">66.46</td>
<td align="center">54.74</td>
<td align="center">57.10</td>
<td align="center">62.71</td>
<td align="center">53.60</td>
</tr>
<tr>
<td align="left">Time information entropy</td>
<td align="center">62.23</td>
<td align="center">71.96</td>
<td align="center">60.86</td>
<td align="center">59.65</td>
<td align="center">68.28</td>
<td align="center">56.77</td>
<td align="center">59.30</td>
<td align="center">64.72</td>
<td align="center">56.03</td>
</tr>
<tr>
<td align="center">Space concentration</td>
<td align="center">68.88</td>
<td align="center">77.61</td>
<td align="center">67.04</td>
<td align="center">66.38</td>
<td align="center">74.70</td>
<td align="center">63.84</td>
<td align="center">65.10</td>
<td align="center">72.07</td>
<td align="center">62.55</td>
</tr>
<tr>
<td align="center">Spatiotemporal diffusion</td>
<td align="center">72.61</td>
<td align="center">81.50</td>
<td align="center">71.33</td>
<td align="center">69.94</td>
<td align="center">79.14</td>
<td align="center">67.93</td>
<td align="center">68.84</td>
<td align="center">76.51</td>
<td align="center">66.38</td>
</tr>
<tr>
<td align="center">Energy concentration index</td>
<td align="center">77.38</td>
<td align="center">86.28</td>
<td align="center">75.40</td>
<td align="center">74.67</td>
<td align="center">83.53</td>
<td align="center">72.23</td>
<td align="center">73.31</td>
<td align="center">81.23</td>
<td align="center">70.90</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
</sec>
<sec sec-type="conclusion" id="s7">
<title>7 Conclusion</title>
<p>To overcome the problems of previous rockburst prediction tasks, this paper constructs quantitative information models of geology and mining based on typical rockburst events. By studying the relationship between the spatial-temporal distribution characteristics of MS data before a rockburst and the main controlling factors of rockbursts, the results show that the distribution features may be different for the same type of MS and rockburst events. Different types of events may show similar distribution features. The quantitative research results on the relationship between the deep learning prediction algorithm performance and prediction vectors show that the rockburst prediction model must accurately distinguish among the various types of rockburst events and also accurately describe the intra-class variance of monitoring data.</p>
<p>Based on the above insights, a novel rockburst risk and type prediction algorithm based on a CNN-GRU model with prototype-based prediction is proposed. The CNN-GRU model consists of an encoder and a decoder, which can produce prediction vectors by fusing implicit and explicit information extracted from the original MS data and prediction indexes. The prototype-based prediction uses cross-entropy loss, vector-prototype contrastive loss, and vector-prototype contrastive loss to automatically control the intra-class variance and inter-class gap of rockburst risk and type prediction vectors. The performance superiority of the proposed algorithm compared with the previous algorithm is verified by the comparison experiment on the data of three mines. The ablation experiment also proves the universality of the proposed prototype-based prediction head in different algorithms for rockburst risk and type prediction.</p>
</sec>
</body>
<back>
<sec sec-type="data-availability" id="s8">
<title>Data availability statement</title>
<p>The original contributions presented in the study are included in the article/supplementary material, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec sec-type="author-contributions" id="s9">
<title>Author contributions</title>
<p>XZ: Writing &#x2013; original draft, Writing &#x2013; review and editing. GL: Writing &#x2013; review and editing, Writing &#x2013; original draft. YC: Writing &#x2013; original draft, Writing &#x2013; review and editing. HW: Writing &#x2013; review and editing, Writing &#x2013; original draft. HZ: Writing &#x2013; original draft, Writing &#x2013; review and editing. HL: Writing &#x2013; original draft, Writing &#x2013; review and editing. WD: Conceptualization, Visualization, Investigation, Software, Validation, Funding acquisition, Formal analysis, Supervision, Writing &#x2013; review and editing, Resources, Data curation, Methodology, Project administration, Writing &#x2013; original draft. XL: Writing &#x2013; original draft, Writing &#x2013; review and editing. XX: Writing &#x2013; original draft, Writing &#x2013; review and editing. YH: Writing &#x2013; review and editing, Writing &#x2013; original draft.</p>
</sec>
<sec sec-type="funding-information" id="s10">
<title>Funding</title>
<p>The author(s) declare that financial support was received for the research and/or publication of this article. The authors would like to acknowledge the support of Chinese Institute of Coal Science, Beijing 100013, China; the Shandong Energy Group Co., LTD, Jinan Shandong 250014, China (No. SNKJ2023A17-R01, No. SNKJ2022BJ03-R28); the Taishan Industrial Experts Program (No. tscx202408130); and the Key Science and Technology Project of the Ministry of Emergency Management of the People&#x27;s Republic of China (2024EMST070702).</p>
</sec>
<sec sec-type="COI-statement" id="s11">
<title>Conflict of interest</title>
<p>Authors XZ, GL, YC, and HW were employed by Shandong Energy Group Co., Ltd.</p>
<p>The remaining authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
<p>The authors declare that this study received funding from Shandong Energy Group. The funder had the following involvement in the study: data collection and analysis, decision to publish, and preparation of the manuscript.</p>
</sec>
<sec sec-type="ai-statement" id="s12">
<title>Generative AI statement</title>
<p>The author(s) declare that no Generative AI was used in the creation of this manuscript.</p>
</sec>
<sec sec-type="disclaimer" id="s13">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Adoko</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Zvarivadza</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>A bayesian approach for predicting rockburst</article-title>. in <conf-name>ARMA US Rock Mechanics/Geomechanics Symposium</conf-name>. <conf-loc>Seattle, Washington</conf-loc> <publisher-name>ARMA</publisher-name>.</citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Askaripour</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Saeidi</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Rouleau</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Mercier-Langevin</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Rockburst in underground excavations: a review of mechanism, classification, and prediction methods</article-title>. <source>Undergr. Space</source> <volume>7</volume> (<issue>4</issue>), <fpage>577</fpage>&#x2013;<lpage>607</lpage>. <pub-id pub-id-type="doi">10.1016/j.undsp.2021.11.008</pub-id>
</citation>
</ref>
<ref id="B4">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Aydan</surname>
<given-names>&#xd6;.</given-names>
</name>
<name>
<surname>Geni&#x15f;</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Akagi</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Kawamoto</surname>
<given-names>T.</given-names>
</name>
</person-group> (<year>2017</year>). <source>Assessment of susceptibility of rock bursting in tunnelling in hard rocks, Modern tunneling science and technology</source>. <publisher-loc>London</publisher-loc>: <publisher-name>Routledge</publisher-name>, <fpage>391</fpage>&#x2013;<lpage>396</lpage>.</citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Basnet</surname>
<given-names>P. M. S.</given-names>
</name>
<name>
<surname>Mahtab</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Jin</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>A comprehensive review of intelligent machine learning based predicting methods in long-term and short-term rock burst prediction</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>142</volume>, <fpage>105434</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2023.105434</pub-id>
</citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Liang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Cao</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>ConvLSTM for predicting short-term spatiotemporal distribution of seismic risk induced by large-scale coal mining</article-title>. <source>Nat. Resour. Res.</source> <volume>32</volume> (<issue>3</issue>), <fpage>1459</fpage>&#x2013;<lpage>1479</lpage>. <pub-id pub-id-type="doi">10.1007/s11053-023-10193-5</pub-id>
</citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Yin</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Gao</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2024</year>). <article-title>Microseismic data-driven short-term rockburst evaluation in underground engineering with strategic data augmentation and extremely randomized forest</article-title>. <source>Mathematics</source> <volume>12</volume> (<issue>22</issue>), <fpage>3502</fpage>. <pub-id pub-id-type="doi">10.3390/math12223502</pub-id>
</citation>
</ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cheng</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Qiao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Experimental investigation on the influence of a single structural plane on rockburst</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>132</volume>, <fpage>104914</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2022.104914</pub-id>
</citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Di</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2023a</year>). <article-title>Predicting microseismic, acoustic emission and electromagnetic radiation data using neural networks</article-title>. <source>J. Rock Mech. Geotechnical Eng.</source> <volume>16</volume>, <fpage>616</fpage>&#x2013;<lpage>629</lpage>. <pub-id pub-id-type="doi">10.1016/j.jrmge.2023.05.012</pub-id>
</citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Di</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2023b</year>). <article-title>Comprehensive early warning method of microseismic, acoustic emission, and electromagnetic radiation signals of rock burst based on deep learning</article-title>. <source>Int. J. Rock Mech. Min. Sci.</source> <volume>170</volume>, <fpage>105519</fpage>. <pub-id pub-id-type="doi">10.1016/j.ijrmms.2023.105519</pub-id>
</citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Dong</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Shu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Tang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Yan</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Microseismic event waveform classification using CNN-based transfer learning models</article-title>. <source>Int. J. Min. Sci. Technol.</source> <volume>33</volume>, <fpage>1203</fpage>&#x2013;<lpage>1216</lpage>. <pub-id pub-id-type="doi">10.1016/j.ijmst.2023.09.003</pub-id>
</citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Qiao</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>A review of rockburst: experiments, theories, and simulations</article-title>. <source>J. Rock Mech. Geotechnical Eng.</source> <volume>15</volume> (<issue>5</issue>), <fpage>1312</fpage>&#x2013;<lpage>1353</lpage>. <pub-id pub-id-type="doi">10.1016/j.jrmge.2022.07.014</pub-id>
</citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Ren</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Rockburst mechanism research and its control</article-title>. <source>Int. J. Min. Sci. Technol.</source> <volume>28</volume> (<issue>5</issue>), <fpage>829</fpage>&#x2013;<lpage>837</lpage>. <pub-id pub-id-type="doi">10.1016/j.ijmst.2018.09.002</pub-id>
</citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hu</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Feng</surname>
<given-names>X.-T.</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>Z.-B.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Niu</surname>
<given-names>W.-J.</given-names>
</name>
<name>
<surname>Bi</surname>
<given-names>X.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>Rockburst time warning method with blasting cycle as the unit based on microseismic information time series: a case study</article-title>. <source>Bull. Eng. Geol. Environ.</source> <volume>82</volume> (<issue>4</issue>), <fpage>121</fpage>. <pub-id pub-id-type="doi">10.1007/s10064-023-03141-3</pub-id>
</citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ji</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Xie</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Song</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Investigate contribution of multi-microseismic data to rockburst risk prediction using support vector machine with genetic algorithm</article-title>. <source>IEEE Access</source> <volume>8</volume>, <fpage>58817</fpage>&#x2013;<lpage>58828</lpage>. <pub-id pub-id-type="doi">10.1109/access.2020.2982366</pub-id>
</citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jin</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Basnet</surname>
<given-names>P. M. S.</given-names>
</name>
<name>
<surname>Mahtab</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Microseismicity-based short-term rockburst prediction using non-linear support vector machine</article-title>. <source>Acta Geophys.</source> <volume>70</volume> (<issue>4</issue>), <fpage>1717</fpage>&#x2013;<lpage>1736</lpage>. <pub-id pub-id-type="doi">10.1007/s11600-022-00817-4</pub-id>
</citation>
</ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Jinqiang</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Basnet</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Mahtab</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Review of machine learning and deep learning application in mine microseismic event classification</article-title>. <source>Mining Mineral Deposits</source> <volume>15</volume>, <fpage>19</fpage>&#x2013;<lpage>26</lpage>. <pub-id pub-id-type="doi">10.33271/mining15.01.019</pub-id>
</citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Yuyang</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Characteristics of microseismic waveforms induced by hydraulic fracturing in coal seam for coal rock dynamic disasters prevention</article-title>. <source>Saf. Sci.</source> <volume>115</volume>, <fpage>188</fpage>&#x2013;<lpage>198</lpage>. <pub-id pub-id-type="doi">10.1016/j.ssci.2019.01.024</pub-id>
</citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Zare Naghadehi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Jimenez</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Evaluating short-term rock burst damage in underground mines using a systems approach</article-title>. <source>Int. J. Min. Reclam. Environ.</source> <volume>34</volume> (<issue>8</issue>), <fpage>531</fpage>&#x2013;<lpage>561</lpage>. <pub-id pub-id-type="doi">10.1080/17480930.2019.1657654</pub-id>
</citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liang</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Sari</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>McKinnon</surname>
<given-names>S. D.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Short-term rockburst risk prediction using ensemble learning methods</article-title>. <source>Nat. Hazards</source> <volume>104</volume>, <fpage>1923</fpage>&#x2013;<lpage>1946</lpage>. <pub-id pub-id-type="doi">10.1007/s11069-020-04255-7</pub-id>
</citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liang</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Sari</surname>
<given-names>Y. A.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>McKinnon</surname>
<given-names>S. D.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Probability estimates of short-term rockburst risk with ensemble classifiers</article-title>. <source>Rock Mech. Rock Eng.</source> <volume>54</volume>, <fpage>1799</fpage>&#x2013;<lpage>1814</lpage>. <pub-id pub-id-type="doi">10.1007/s00603-021-02369-3</pub-id>
</citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liang</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>G.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>A review of long-term and short-term rockburst risk evaluations in deep hard rock</article-title>. <source>J. Rock Mech. Eng.</source> <volume>41</volume>, <fpage>19</fpage>&#x2013;<lpage>39</lpage>. <pub-id pub-id-type="doi">10.13722/j.cnki.jrme.2021.0165</pub-id>
</citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ma</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Jun</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Research on prediction of rockburst microseismic parameters based on CNN-LSTM hybrid model</article-title>. <source>IOP Conf. Ser. Earth Environ. Sci.</source> <volume>861</volume>, <fpage>052097</fpage>. <pub-id pub-id-type="doi">10.1088/1755-1315/861/5/052097</pub-id>
</citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Manouchehrian</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Cai</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Numerical modeling of rockburst near fault zones in deep tunnels</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>80</volume>, <fpage>164</fpage>&#x2013;<lpage>180</lpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2018.06.015</pub-id>
</citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Pu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Apel</surname>
<given-names>D. B.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Mitri</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Machine learning methods for rockburst prediction-state-of-the-art review</article-title>. <source>Int. J. Min. Sci. Technol.</source> <volume>29</volume> (<issue>4</issue>), <fpage>565</fpage>&#x2013;<lpage>570</lpage>. <pub-id pub-id-type="doi">10.1016/j.ijmst.2019.06.009</pub-id>
</citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Apel</surname>
<given-names>D. B.</given-names>
</name>
<name>
<surname>Pu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Hall</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Wei</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Sepehri</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Numerical modeling for rockbursts: a state-of-the-art review</article-title>. <source>J. Rock Mech. Geotechnical Eng.</source> <volume>13</volume> (<issue>2</issue>), <fpage>457</fpage>&#x2013;<lpage>478</lpage>. <pub-id pub-id-type="doi">10.1016/j.jrmge.2020.09.011</pub-id>
</citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Xue</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Song</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>C.</given-names>
</name>
<etal/>
</person-group> (<year>2023</year>). <article-title>A method to predict rockburst using temporal trend test and its application</article-title>. <source>J. Rock Mech. Geotechnical Eng.</source> <volume>16</volume>, <fpage>909</fpage>&#x2013;<lpage>923</lpage>. <pub-id pub-id-type="doi">10.1016/j.jrmge.2023.07.017</pub-id>
</citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Tan</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Analytical estimation of stress distribution in interbedded layers and its implication to rockburst in strong layer</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>81</volume>, <fpage>289</fpage>&#x2013;<lpage>295</lpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2018.07.007</pub-id>
</citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yin</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>X.</given-names>
</name>
<etal/>
</person-group> (<year>2024b</year>). <article-title>Probabilistic assessment of rockburst risk in TBM-excavated tunnels with multi-source data fusion</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>152</volume>, <fpage>105915</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2024.105915</pub-id>
</citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yin</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2021a</year>). <article-title>Real-time prediction of rockburst intensity using an integrated CNN-Adam-BO algorithm based on microseismic data and its engineering application</article-title>. <source>Tunn. Undergr. Space Technol.</source> <volume>117</volume>, <fpage>104133</fpage>. <pub-id pub-id-type="doi">10.1016/j.tust.2021.104133</pub-id>
</citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yin</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Lei</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Lei</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2024a</year>). <article-title>Hybrid deep learning-based identification of microseismic events in TBM tunnelling</article-title>. <source>Measurement</source> <volume>238</volume>, <fpage>115381</fpage>. <pub-id pub-id-type="doi">10.1016/j.measurement.2024.115381</pub-id>
</citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yin</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2021b</year>). <article-title>A novel tree-based algorithm for real-time prediction of rockburst risk using field microseismic monitoring</article-title>. <source>Environ. Earth Sci.</source> <volume>80</volume>, <fpage>504</fpage>&#x2013;<lpage>519</lpage>. <pub-id pub-id-type="doi">10.1007/s12665-021-09802-4</pub-id>
</citation>
</ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yin</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Pan</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>X.</given-names>
</name>
</person-group> (<year>2021c</year>). <article-title>Strength of stacking technique of ensemble learning in rockburst prediction with imbalanced data: comparison of eight single and ensemble models</article-title>. <source>Nat. Resour. Res.</source> <volume>30</volume>, <fpage>1795</fpage>&#x2013;<lpage>1815</lpage>. <pub-id pub-id-type="doi">10.1007/s11053-020-09787-0</pub-id>
</citation>
</ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Zeng</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Fang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Ma</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Yao</surname>
<given-names>Z.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Time series prediction of microseismic multi-parameter related to rockburst based on deep learning</article-title>. <source>Rock Mech. Rock Eng.</source> <volume>54</volume> (<issue>12</issue>), <fpage>6299</fpage>&#x2013;<lpage>6321</lpage>. <pub-id pub-id-type="doi">10.1007/s00603-021-02614-9</pub-id>
</citation>
</ref>
</ref-list>
</back>
</article>