<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article article-type="research-article" dtd-version="2.3" xml:lang="EN" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Energy Res.</journal-id>
<journal-title>Frontiers in Energy Research</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Energy Res.</abbrev-journal-title>
<issn pub-type="epub">2296-598X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="publisher-id">763977</article-id>
<article-id pub-id-type="doi">10.3389/fenrg.2021.763977</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Energy Research</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Municipal Solid Waste Forecasting in China Based on Machine Learning Models</article-title>
<alt-title alt-title-type="left-running-head">Yang et&#x20;al.</alt-title>
<alt-title alt-title-type="right-running-head">Solid Waste Forecasting Machine Learning</alt-title>
</title-group>
<contrib-group>
<contrib contrib-type="author">
<name>
<surname>Yang</surname>
<given-names>Liping</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1454466/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Zhao</surname>
<given-names>Yigang</given-names>
</name>
<xref ref-type="aff" rid="aff3">
<sup>3</sup>
</xref>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Niu</surname>
<given-names>Xiaxia</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1454483/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Song</surname>
<given-names>Zisheng</given-names>
</name>
<xref ref-type="aff" rid="aff4">
<sup>4</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1512688/overview"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Gao</surname>
<given-names>Qingxian</given-names>
</name>
<xref ref-type="aff" rid="aff5">
<sup>5</sup>
</xref>
</contrib>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Wu</surname>
<given-names>Jun</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="corresp" rid="c001">&#x2a;</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1365085/overview"/>
</contrib>
</contrib-group>
<aff id="aff1">
<label>
<sup>1</sup>
</label>School of Economics and Management, Beijing University of Chemical Technology, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<aff id="aff2">
<label>
<sup>2</sup>
</label>School of Management, University of Science and Technology of China, <addr-line>Anhui</addr-line>, <country>China</country>
</aff>
<aff id="aff3">
<label>
<sup>3</sup>
</label>Beijing Institute of Petrochemical Technology, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<aff id="aff4">
<label>
<sup>4</sup>
</label>Department of International Exchange and Cooperation, Beijing University of Chemical Technology, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<aff id="aff5">
<label>
<sup>5</sup>
</label>Chinese Research Academy of Environmental Sciences, <addr-line>Beijing</addr-line>, <country>China</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>
<bold>Edited by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/79597/overview">Xun Zhang</ext-link>, Academy of Mathematics and Systems Science (CAS), China</p>
</fn>
<fn fn-type="edited-by">
<p>
<bold>Reviewed by:</bold> <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/615262/overview">Kenneth E. Okedu</ext-link>, National University of Science and Technology (Muscat),&#x20;Oman</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/950488/overview">Bo Liu</ext-link>, King Fahd University of Petroleum and Minerals, Saudi Arabia</p>
</fn>
<corresp id="c001">&#x2a;Correspondence: Jun Wu, <email>wujun@mail.buct.edu.cn</email>
</corresp>
<fn fn-type="other">
<p>This article was submitted to Smart Grids, a section of the journal Frontiers in Energy Research</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>08</day>
<month>11</month>
<year>2021</year>
</pub-date>
<pub-date pub-type="collection">
<year>2021</year>
</pub-date>
<volume>9</volume>
<elocation-id>763977</elocation-id>
<history>
<date date-type="received">
<day>24</day>
<month>08</month>
<year>2021</year>
</date>
<date date-type="accepted">
<day>19</day>
<month>10</month>
<year>2021</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2021 Yang, Zhao, Niu, Song, Gao and Wu.</copyright-statement>
<copyright-year>2021</copyright-year>
<copyright-holder>Yang, Zhao, Niu, Song, Gao and Wu</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these&#x20;terms.</p>
</license>
</permissions>
<abstract>
<p>As the largest producing country of municipal solid waste (MSW) around the world, China is always challenged by a lower utilization rate of MSW due to a lack of a smart MSW forecasting strategy. This paper mainly aims to construct an effective MSW prediction model to handle this problem by using machine learning techniques. Based on the empirical analysis of provincial panel data from 2008 to 2019 in China, we find that the Deep Neural Network (DNN) model performs best among all machine learning models. Additionally, we introduce the <italic>SHapley Additive exPlanation</italic> (SHAP) method to unravel the correlation between MSW production and socioeconomic features (e.g., total regional GDP, population density). We also find the increase of urban population and agglomeration of wholesales and retails industries can positively promote the production of MSW in regions of high economic development, and vice versa. These results can be of help in the planning, design, and implementation of solid waste management system in China.</p>
</abstract>
<kwd-group>
<kwd>municipal solid waste</kwd>
<kwd>influencing factors</kwd>
<kwd>machine learning</kwd>
<kwd>deep learning</kwd>
<kwd>SHAP value</kwd>
</kwd-group>
</article-meta>
</front>
<body>
<sec id="s1">
<title>Introduction</title>
<p>Over the past decade, the urban population in China has reached up to 900 million residents with an urbanization rate of over 60% (<xref ref-type="bibr" rid="B34">NBSC, 2021</xref>), which significantly challenges the existing urban sources (e.g., water, air, and energy) related to residents&#x2019; life quality (<xref ref-type="bibr" rid="B17">Hoornweg and Bhada-Tata, 2012</xref>). The municipal solid waste (MSW), as renewable energy, is considered an essential part of the Waste-to-Energy (WtE) system (<xref ref-type="bibr" rid="B37">Ouda et&#x20;al., 2013</xref>; <xref ref-type="bibr" rid="B22">Kuznetsova et&#x20;al., 2019</xref>; <xref ref-type="bibr" rid="B31">Mukherjee et&#x20;al., 2020</xref>). It is reported that the production of MSW in China was around 242 million tons in 2020 compared with that of 8.17 million tons in 2008 (<xref ref-type="bibr" rid="B33">NBSC, 2020</xref>). In other words, the efficient management of municipal solid waste is becoming an important concern for urban sustainability governance. However, the utilization efficiency of MSW was merely about 45% in China, which was much lower than that in other advanced countries, such as over 80% in Japan (<xref ref-type="bibr" rid="B12">Ding et&#x20;al., 2021</xref>). Therefore, how to increase the utilization efficiency of MSW would impact both central and local governments in China to promote urban sustainable development (<xref ref-type="bibr" rid="B16">He and Lin, 2019</xref>).</p>
<p>In general, an integrated decision-support methodology for waste-to-energy management systems (WtEMS) design is mainly composed of three modules: 1) the waste modeling and prediction, 2) optimization of WtEMS, and 3) a multi-dimensional assessment, as shown in <xref ref-type="fig" rid="F1">Figure&#x20;1</xref> (<xref ref-type="bibr" rid="B22">Kuznetsova et&#x20;al., 2019</xref>). Among these three modules, waste modeling and its prediction of MSW play a fundamental role in effectively conducting urban planning and energy management. Many international scholars have carried out extensive studies on this module by using group comparisons, time series analysis, and system dynamics (<xref ref-type="bibr" rid="B4">Beigl et&#x20;al., 2008</xref>). Recently, with the popularity of machine learning (ML) methods, alternative methods were put forward to forecast the quantity of generated municipal solid waste effectively (<xref ref-type="bibr" rid="B14">Guo et&#x20;al., 2021</xref>). For instance, based on the example of Suzhou (<xref ref-type="bibr" rid="B36">Niu et&#x20;al., 2021</xref>), constructed the long short-term memory (LSTM) neural network, autoregressive integrated moving average (ARIMA), and traditional neural network to predict the MSW production. They found that the LSTM played a vital role in predicting MSW production but did not reveal the correlation between the production of MSW and socio-economic variables. <xref ref-type="bibr" rid="B35">Nguyen et&#x20;al. (2021)</xref> selected residential areas in Vietnam as a case of study and figured out that both the random forest (RF) and the k-nearest neighbor (KNN) approaches performed effectively in predicting the amount of urban waste. <xref ref-type="bibr" rid="B5">Birgen et&#x20;al. (2021)</xref> developed a Gaussian Processes Regression (GPR) method to predict the daily lower heating value of MSW by combining the historical data of a WtE plant and the weather and calendar data. In addition, other ML methods, such as the support vector machine (SVM) (<xref ref-type="bibr" rid="B21">Kumar et&#x20;al., 2018</xref>) and decision tree (<xref ref-type="bibr" rid="B20">Kannangara et&#x20;al., 2018</xref>) have also been employed to predict the MSW production.</p>
<fig id="F1" position="float">
<label>FIGURE 1</label>
<caption>
<p>Integrated decision support method for WtEMS design: methodology flowchart.</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g001.tif"/>
</fig>
<p>Similar to other energy forecasting research topics (e.g., crude oil prices, gas consumption), MSW production is also was highly influenced by various socio-economic factors (<xref ref-type="bibr" rid="B40">Zhang et&#x20;al., 2009</xref>; <xref ref-type="bibr" rid="B24">Liang et&#x20;al., 2019</xref>; <xref ref-type="bibr" rid="B18">Huang et&#x20;al., 2021a</xref>). However, previous studies neither revealed the correlation between each factor and MSW production nor identified their interaction in different socio-economic circumstances (<xref ref-type="bibr" rid="B20">Kannangara et&#x20;al., 2018</xref>; <xref ref-type="bibr" rid="B36">Niu et&#x20;al., 2021</xref>; <xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021</xref>). In the context of China, existing studies scarcely discussed the performances and applications of different ML methods in predicting MSW. Therefore, this paper mainly aimed to construct a prediction model by using machine learning models by using provincial panel data of 2008&#x2013;2019 in China. Besides, it also discussed the comparison of the performances of six different ML models in predicting China&#x2019;s municipal solid waste generation. Considering that data input form and model hyperparameters have a great influence on prediction results, we tested different preprocessing strategies to ensure robust estimation and prediction of the ML model. Finally, this paper provided some potential implications for both policy-makers and other industry stakeholders in terms of convincing evidence concluded from the ML prediction&#x20;model.</p>
<p>The initial contributions of this paper are threefold. First, it emphasized the good performance of machine learning approaches in predicting MSW production and extended the existing literature to construct a prediction model by comparing six supervised learning algorithms. These models varied from linear, non-linear to ensemble methods and artificial neural network methods, including a body of discussions on data preprocessing, resampling, model training, testing, and interpretation steps. Therefore, the constructed prediction model of MSW would theoretically shed light on other similar research related to prediction issues in the future. Second, this paper estimated the impacts of diverse socio-economic factors on MSW production, such as the regional economic development level (e.g., regional GDP, population density, per capita disposable income), industrial structure (e.g., wholesale and retail values added), and waste generation characteristics. Third, to improve the interpretations of ML models, this paper employed the <italic>SHapley Additive exPlanation</italic> (SHAP) approach and visualized the SHAP value of each explanatory variable. This technique would also provide good evidence to explain the outcomes of ML models for other researchers in the future.</p>
<p>The remaining sections of this paper are organized as follows: <italic>Materials and Methods</italic> describes the models adopted in this paper and the process of data acquisition. <italic>Results</italic> reports the results of comparison among six ML models, <italic>via</italic> presenting the predictive capability and SHAP analysis. <italic>Conclusion</italic> provides conclusions and some implications.</p>
</sec>
<sec sec-type="materials|methods" id="s2">
<title>Materials and Methods</title>
<p>
<xref ref-type="fig" rid="F2">Figure&#x20;2</xref> outlines the main steps of the methodology used in this study. In this paper, we first preprocessed the original database and selected critical variables for MSW prediction. Second, this paper focused on comparing with six ML models, including the multiple linear regression (MLR), support vector regression (SVR), Random Forest, extreme gradient boosting (XGBoost), k-nearest neighbor, and deep neural network (DNN). Thirdly, three evaluation metrics are used to compare the prediction performance of each algorithm. Finally, the SHAP method is employed to analyze and discuss the output.</p>
<fig id="F2" position="float">
<label>FIGURE 2</label>
<caption>
<p>Procedures of methodology.</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g002.tif"/>
</fig>
<sec id="s2-1">
<title>ML-Based Models and Applications for Waste Prediction</title>
<sec id="s2-1-1">
<title>The Multiple Linear Regression Liner Model</title>
<p>The multiple linear regression is a commonly used ML method to estimate the marginal effects of independent variables (or called feature vector in machine learning techniques) on the dependent variable. It is widely applied to waste prediction of desirable explanatory power in different regions and countries (<xref ref-type="bibr" rid="B4">Beigl et&#x20;al., 2008</xref>). In China, this approach is also employed to predict the MSW production in &#x201c;Calculation and Prediction Method of Municipal Solid Waste Production (CJ/T 106-1999)&#x201d;, which is the official guide compiled by the Ministry of Construction, China.</p>
<p>The model can be expressed as <xref ref-type="disp-formula" rid="e1">Eq. 1</xref>:<disp-formula id="e1">
<mml:math id="m1">
<mml:mrow>
<mml:mi>Y</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:msub>
<mml:mi>&#x3b2;</mml:mi>
<mml:mn>0</mml:mn>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>&#x3b2;</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:msub>
<mml:mi>X</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>&#x2b;</mml:mo>
<mml:msub>
<mml:mi>&#x3b2;</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
<mml:msub>
<mml:mi>X</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
<mml:mo>&#x2b;</mml:mo>
<mml:mi mathvariant="italic">&#x3f5;</mml:mi>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(1)</label>
</disp-formula>where <inline-formula id="inf1">
<mml:math id="m2">
<mml:mi>Y</mml:mi>
</mml:math>
</inline-formula> is MSW generation in this paper, <inline-formula id="inf2">
<mml:math id="m3">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3b2;</mml:mi>
<mml:mn>0</mml:mn>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> denotes regression constant,<inline-formula id="inf3">
<mml:math id="m4">
<mml:mrow>
<mml:mtext>&#xa0;</mml:mtext>
<mml:msub>
<mml:mi>&#x3b2;</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>&#x223c;</mml:mo>
<mml:msub>
<mml:mi>&#x3b2;</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> are regression coefficients, <inline-formula id="inf4">
<mml:math id="m5">
<mml:mrow>
<mml:msub>
<mml:mi>X</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>&#x223c;</mml:mo>
<mml:msub>
<mml:mi>X</mml:mi>
<mml:mi>k</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> are explanatory variables, <inline-formula id="inf5">
<mml:math id="m6">
<mml:mi mathvariant="italic">&#x3f5;</mml:mi>
</mml:math>
</inline-formula> marks the regression residuals.</p>
<p>Usually, MLR uses the ordinary least squares (OLS) method to estimate the parameters that can achieve the lowest sum-of-squared errors between the observed and predicted responses. Under the OLS estimation, MLR&#x2019;s results could be easily interpreted. However, some drawbacks have to be considered in MLR. For instance, the multicollinearity among the predictors can result in estimation errors, as well as the omitted variables could induce a biased estimation. In this paper, we mainly concentrated on the performance of each ML model and considered the variables selection based on earlier studies (<xref ref-type="bibr" rid="B20">Kannangara et&#x20;al., 2018</xref>; <xref ref-type="bibr" rid="B32">Namlis and Komilis, 2019</xref>; <xref ref-type="bibr" rid="B36">Niu et&#x20;al., 2021</xref>; <xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021</xref>). The multicollinearity and omitted variables problems are not our concerns.</p>
</sec>
<sec id="s2-1-2">
<title>Support Vector Regression</title>
<p>SVM was originally used to deal with pattern recognition problems, and recently extended to estimate regression models due to its properties of the sparse solution and good generalization (<xref ref-type="bibr" rid="B11">Demir and Bruzzone, 2014</xref>). By introducing an <inline-formula id="inf6">
<mml:math id="m7">
<mml:mi>&#x3b5;</mml:mi>
</mml:math>
</inline-formula>-tube to reformulate the optimization problem, the SVM model could be transformed to an SVR model and finds the optimal approximation of the continuous-valued function while balancing the complexity and prediction error of the prediction model (<xref ref-type="bibr" rid="B19">Huang et&#x20;al., 2021b</xref>). In addition, the accuracy of an SVR model heavily relies on three parameters: a penalty parameter (<inline-formula id="inf7">
<mml:math id="m8">
<mml:mi>C</mml:mi>
</mml:math>
</inline-formula>), the kernel width (<inline-formula id="inf8">
<mml:math id="m9">
<mml:mi>&#x3b3;</mml:mi>
</mml:math>
</inline-formula>) and the precision parameter (<inline-formula id="inf9">
<mml:math id="m10">
<mml:mi>&#x3b5;</mml:mi>
</mml:math>
</inline-formula>) (<xref ref-type="bibr" rid="B1">Abbasi and El Hanandeh, 2016</xref>; <xref ref-type="bibr" rid="B23">Li et&#x20;al., 2021</xref>). Specifically, the smaller <inline-formula id="inf10">
<mml:math id="m11">
<mml:mi>C</mml:mi>
</mml:math>
</inline-formula> is, the smaller the fitting error and the weaker the generalization ability would be. The larger <inline-formula id="inf11">
<mml:math id="m12">
<mml:mi>&#x3b3;</mml:mi>
</mml:math>
</inline-formula> is, the more support vectors; and vice versa. <inline-formula id="inf12">
<mml:math id="m13">
<mml:mi>&#x3b5;</mml:mi>
</mml:math>
</inline-formula> is a precision parameter representing the tube&#x2019;s radius located around the regression function. In other words, the choice of <inline-formula id="inf13">
<mml:math id="m14">
<mml:mi>&#x3b5;</mml:mi>
</mml:math>
</inline-formula> donates the magnitude of errors that can be neglected. Since the above three parameters are critical to the adaptability of the model, we will tune them using a grid optimization approach in <italic>Results</italic> to optimize the SVR&#x20;model.</p>
<p>A great body of literature has discussed the SVR and SVM models in predicting the generation of MSW. For example (<xref ref-type="bibr" rid="B1">Abbasi and El Hanandeh, 2016</xref>), adjusted the hyper-parameters of SVR by combining the grid search method and applying the model with the optimal parameters to the monthly prediction of MSW in Logan City, Australia. They found that SVR can effectively reduce the mean absolute error (<italic>MAE</italic>) and root-mean-square error (<italic>RMSE</italic>), and improve prediction performance (<italic>R-square</italic>). Besides (<xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021</xref>), applied SVM to the prediction of MSW production in Vietnam with an <italic>MAE</italic> of 131.07, which confirmed that the SVM model performed a better prediction. <xref ref-type="bibr" rid="B21">Kumar et&#x20;al. (2018)</xref> applied it to the prediction of the production rate of plastic waste, and found that the prediction result of SVM (<inline-formula id="inf14">
<mml:math id="m15">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>&#x3d;0.74) is better than RF (<inline-formula id="inf15">
<mml:math id="m16">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>&#x3d;0.66) and lower than artificial neural network (ANN) (<inline-formula id="inf16">
<mml:math id="m17">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>&#x3d;0.75). <xref ref-type="bibr" rid="B30">Mehrdad et&#x20;al. (2021)</xref> argued that SVM was superior to both the adaptive neuro-fuzzy inference system and artificial neural network models in predicting methane generation.</p>
</sec>
<sec id="s2-1-3">
<title>Random Forest</title>
<p>Random Forest is an evolution of Bagging which aims to reduce the variance of a statistical model, simulates the variability of data through the random extraction of bootstrap samples from a single training set and aggregates predictions on a new record (see <xref ref-type="bibr" rid="B6">Breiman, 1996</xref>). It performs a more stable and better prediction of explained variables than other machine learning models (<xref ref-type="bibr" rid="B19">Huang et&#x20;al., 2021b</xref>). Generally, the RF algorithm implementation can be expressed as follows:<list list-type="simple">
<list-item>
<p>1) Bagging is used to randomly generate sample subsets;</p>
</list-item>
<list-item>
<p>2) Use the idea of random subspace by randomly extracting features, splitting nodes, and building a regression sub-decision tree;</p>
</list-item>
<list-item>
<p>3) Repeat the above steps to construct <inline-formula id="inf17">
<mml:math id="m18">
<mml:mi>T</mml:mi>
</mml:math>
</inline-formula> (the number of decision trees) regression decision subtrees to form a random forest;</p>
</list-item>
<list-item>
<p>4) Take the predicted values of <inline-formula id="inf18">
<mml:math id="m19">
<mml:mi>T</mml:mi>
</mml:math>
</inline-formula> sub-decision trees and take the mean as the final prediction result.</p>
</list-item>
</list>
</p>
<p>The RF model was widely used in the prediction of waste. <xref ref-type="bibr" rid="B21">Kumar et&#x20;al. (2018)</xref> used RF for the prediction of plastic waste generation rate that showed an R-square of 0.66. The size of the random forest, that is, the number of decision trees (<inline-formula id="inf19">
<mml:math id="m20">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>) and the number of features tried in each segmentation (<inline-formula id="inf20">
<mml:math id="m21">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>f</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>) have a significant impact on the predictive ability of the RF model (<xref ref-type="bibr" rid="B15">Hariharan, 2021</xref>). When <inline-formula id="inf21">
<mml:math id="m22">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> exceed a certain value, the prediction performance of the model converges. In this case, increasing the number of decision trees will not improve the model performance, but will result in model redundancy. In addition, using a smaller number of<inline-formula id="inf22">
<mml:math id="m23">
<mml:mrow>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> reduces the similarity in the forest, but also reduces the complexity and strength of the model. Conversely, the increase in <inline-formula id="inf23">
<mml:math id="m24">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> can make each tree more powerful, but also increase the correlation between the trees. Therefore, in the following section, we will optimize these two hyper-parameters to acquire better results.</p>
</sec>
<sec id="s2-1-4">
<title>Extreme Gradient Boosting</title>
<p>XGBoost algorithm, proposed in 2016, is a relatively new approach (<xref ref-type="bibr" rid="B8">Chen and Guestrin, 2016</xref>). Different from RF model using bagging integration method, XGBoost model is an integration tree model using boosting method to integrate classification and regression tree (CART). It has the advantages of fast training speed and high prediction accuracy. The result of XGBoost is the sum of prediction scores of all CARTs (<xref ref-type="bibr" rid="B8">Chen and Guestrin, 2016</xref>) as formed in <xref ref-type="disp-formula" rid="e2">Eq. 2</xref>:<disp-formula id="e2">
<mml:math id="m25">
<mml:mrow>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:munderover>
<mml:mstyle displaystyle="true">
<mml:mo>&#x2211;</mml:mo>
</mml:mstyle>
<mml:mrow>
<mml:mi>n</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>N</mml:mi>
</mml:munderover>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mi>m</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mi>X</mml:mi>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(2)</label>
</disp-formula>where <inline-formula id="inf24">
<mml:math id="m26">
<mml:mi>N</mml:mi>
</mml:math>
</inline-formula> represents the number of trees in the model, <inline-formula id="inf25">
<mml:math id="m27">
<mml:mrow>
<mml:mtext>&#xa0;</mml:mtext>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mi>m</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> represents each CART tree and <inline-formula id="inf26">
<mml:math id="m28">
<mml:mrow>
<mml:mover accent="true">
<mml:mi>y</mml:mi>
<mml:mo>&#x5e;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula> is predicted result.</p>
<p>Since its introduction, the XGBoost model has been widely used in the prediction of oil price (<xref ref-type="bibr" rid="B9">Costa et&#x20;al., 2021</xref>) and energy usage (<xref ref-type="bibr" rid="B13">Feng et&#x20;al., 2021</xref>). However, up to date, XGBoost model has not been applied to the research of MSW generation prediction. Similar to RF, the number of integrated CARTs (<inline-formula id="inf27">
<mml:math id="m29">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>) in XGBoost has a great influence on the prediction performance. Therefore, in order to increase the model&#x2019;s performance in predicting the MSW generation, it is necessary to optimize this hyper-parameter. In <italic>Results</italic>, we also use the grid search method to confirm the different combinations of these two parameters to obtain the optimal model structure.</p>
</sec>
<sec id="s2-1-5">
<title>K-Nearest Neighbor</title>
<p>KNN algorithm is a non-parametric learning method first proposed by Cover and Hart (<xref ref-type="bibr" rid="B10">Cover and Hart, 1967</xref>). Since its introduction, it has been widely used in regression and classification due to its simple and intuitive mathematical form (<xref ref-type="bibr" rid="B38">Wu et&#x20;al., 2008</xref>). It is essentially a supervised learning technique that <italic>via</italic> the clustering algorithm classify the similarity between the test sample and <italic>K</italic> nearest training samples (<xref ref-type="bibr" rid="B41">Zheng et&#x20;al., 2020</xref>). Here, <italic>K</italic> is a user-defined number, normally an odd number, and the similarity is measured by the commonly used Euclidean distance. The test sample is classified based on the most frequent classification among the training samples. The mean value of the <italic>K</italic> nearest training samples is regarded as the predicted value. The mathematical measurement of Euclidean distance is expressed in <xref ref-type="disp-formula" rid="e3">Eq. 3</xref>:<disp-formula id="e3">
<mml:math id="m30">
<mml:mrow>
<mml:mi>d</mml:mi>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>y</mml:mi>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x3d;</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mn>1</mml:mn>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x2b;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mn>2</mml:mn>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x2b;</mml:mo>
<mml:mo>&#x2026;</mml:mo>
<mml:mo>&#x2b;</mml:mo>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>n</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:msqrt>
<mml:mo>&#x3d;</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:munderover>
<mml:mstyle displaystyle="true">
<mml:mo>&#x2211;</mml:mo>
</mml:mstyle>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:munderover>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:msqrt>
</mml:mrow>
</mml:math>
<label>(3)</label>
</disp-formula>
</p>
<p>One drawback of KNN approach is the pre-selected number of K, a hyperparameter, because it would greatly influence the numbers of nearest samples (<xref ref-type="bibr" rid="B38">Wu et&#x20;al., 2008</xref>; <xref ref-type="bibr" rid="B41">Zheng et&#x20;al., 2020</xref>). In the following section, we first limit K to positive integers between 1 and 30, and then cross-verify them on a 10-fold sample to avoid this drawback.</p>
<p>Several studies applied the KNN approach into the prediction of MSW. For example, (<xref ref-type="bibr" rid="B1">Abbasi and El Hanandeh, 2016</xref>) first attempt to evaluate the ability of KNN to forecast MSW generation. They concluded that KNN can give good prediction performance and may be applied to establish the forecasting models that could provide accurate and reliable MSW generation prediction. <xref ref-type="bibr" rid="B35">Nguyen et&#x20;al. (2021)</xref> predicted the MSW production in Vietnam and the R-square was over 0.96, which indicated that more than 96% of MSW production would be explained by the KNN&#x20;model.</p>
</sec>
<sec id="s2-1-6">
<title>Artificial Neural Network</title>
<p>The ANN model is a computational system composed of multiple layers of neurons (input-hidden-output) (<xref ref-type="bibr" rid="B2">Al-Dahidi et&#x20;al., 2019</xref>). This model is widely used in waste management because of its strong fault-tolerant ability to describe the complex relationship between variables in a multivariate system. (<xref ref-type="bibr" rid="B1">Abbasi and El Hanandeh, 2016</xref>; <xref ref-type="bibr" rid="B30">Mehrdad et&#x20;al., 2021</xref>; <xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021</xref>; <xref ref-type="bibr" rid="B36">Niu et&#x20;al., 2021</xref>). The deep neural network is a branch of ANN based on a perceptron model. Indeed, an ANN model with multiple hidden layers is called a DNN since it has to train and process through multiple layers (<xref ref-type="bibr" rid="B26">Liu et&#x20;al., 2017</xref>). The structure of DNN also includes input layer, hidden layer, and output layer. In general, the structure of DNN and ANN is similar, and their training algorithm is not different. However, studies showed that DNN tends to provide better performance and accuracy than conventional ANN models (<xref ref-type="bibr" rid="B39">Yang et&#x20;al., 2021</xref>).</p>
<p>In this paper, a DNN with four layers of structure is constructed, namely the input layer, the first hidden layer, the second hidden layer and the output layer with one neuron. The number of neurons in the hidden layer has a great influence on the prediction performance of DNN. The smaller the number of neurons, the more likely it is to lead to insufficient fitting. On the contrary, an excessive number of neurons may lead to over-fitting. Therefore, selecting the appropriate number of neurons for DNN is also one of the bases to improve the model performance. In this paper, the number of neurons in the first hidden layer (<inline-formula id="inf28">
<mml:math id="m31">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>h</mml:mi>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>) and the number of neurons in the second hidden layer (<inline-formula id="inf29">
<mml:math id="m32">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>h</mml:mi>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>) are optimized to gain better results. Specifically, we first specify the numerical space of the number of neurons, and then test on the train and test samples, taking the optimal result as the optimal network structure.</p>
</sec>
</sec>
<sec id="s2-2">
<title>Data Collection</title>
<p>In this paper, we aim to construct a ML-based prediction model of MSW production that is the predictor in all ML models. However, because there are no relevant statistics of MSW production in China at present, we utilize a proxy indicator of the MSW removal volume (<xref ref-type="bibr" rid="B36">Niu et&#x20;al., 2021</xref>; <xref ref-type="bibr" rid="B32">Namlis and Komilis, 2019</xref>). More specifically, we obtained this annual statistical data for all provinces in mainland China from 2008 to 2019 to support our research.</p>
<p>The input variables of this paper in predicting MSW production are collected from provincial panel databases of the China Statistical Yearbook 2008&#x2013;2019. Nine diverse socio-economic factors on MSW production, such as the regional economic development level (e.g., regional GDP, population density, per capita disposable income), industrial structure (e.g., wholesale and retail values added), and waste generation characteristics are obtained (<xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021</xref>). <xref ref-type="table" rid="T1">Table&#x20;1</xref> reported the variable definition and descriptive statistics. As plotted in <xref ref-type="fig" rid="F3">Figure&#x20;3</xref>, the skewness and kurtosis of each variable existed noticeable differences. To mitigate the influences in predicting the MSW production, we employ three different data preprocessing methods and proceed to explore the model&#x2019;s performance under different circumstances in the following sub-sections.</p>
<table-wrap id="T1" position="float">
<label>TABLE 1</label>
<caption>
<p>Definition of variables and descriptive statistics.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Category</th>
<th align="center">Variable</th>
<th align="center">Description</th>
<th align="center">Mean</th>
<th align="center">Median</th>
<th align="center">Maximum</th>
<th align="center">Minimum</th>
<th align="center">
<inline-formula id="inf30">
<mml:math id="m33">
<mml:mrow>
<mml:mi mathvariant="bold-italic">Std</mml:mi>
<mml:mo>.</mml:mo>
<mml:mi mathvariant="bold-italic">&#xa0;Dec</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</th>
<th align="center">Unit</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Explained variable</td>
<td align="left">
<inline-formula id="inf31">
<mml:math id="m34">
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>W</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Total solid waste collected amount</td>
<td align="center">
<inline-formula id="inf32">
<mml:math id="m35">
<mml:mrow>
<mml:mn>8343.85</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">6125.25</td>
<td align="char" char=".">42951.80</td>
<td align="char" char=".">130.00</td>
<td align="center">
<inline-formula id="inf33">
<mml:math id="m36">
<mml:mrow>
<mml:mn>7767.28</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">10,000 tons</td>
</tr>
<tr>
<td rowspan="9" align="left">Explanatory variables</td>
<td align="left">
<inline-formula id="inf34">
<mml:math id="m37">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>G</mml:mi>
<mml:mi>D</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Total Regional GDP.</td>
<td align="center">
<inline-formula id="inf35">
<mml:math id="m38">
<mml:mrow>
<mml:mn>20265.24</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">14580.35</td>
<td align="char" char=".">107986.90</td>
<td align="char" char=".">398.20</td>
<td align="center">
<inline-formula id="inf36">
<mml:math id="m39">
<mml:mrow>
<mml:mn>18414.77</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">100 million RMB</td>
</tr>
<tr>
<td align="left">
<inline-formula id="inf37">
<mml:math id="m40">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>T</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Value added by transportation, warehousing, and postal services</td>
<td align="center">
<inline-formula id="inf38">
<mml:math id="m41">
<mml:mrow>
<mml:mn>932.68</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">727.80</td>
<td align="char" char=".">3658.00</td>
<td align="char" char=".">20.60</td>
<td align="center">
<inline-formula id="inf39">
<mml:math id="m42">
<mml:mrow>
<mml:mn>746.80</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">100 million RMB</td>
</tr>
<tr>
<td align="left">
<inline-formula id="inf40">
<mml:math id="m43">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>W</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Wholesale and retail value added</td>
<td align="center">
<inline-formula id="inf41">
<mml:math id="m44">
<mml:mrow>
<mml:mn>1955.78</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">1250.85</td>
<td align="char" char=".">11000.20</td>
<td align="char" char=".">23.40</td>
<td align="center">
<inline-formula id="inf42">
<mml:math id="m45">
<mml:mrow>
<mml:mn>2097.82</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">100 million RMB</td>
</tr>
<tr>
<td align="left">
<inline-formula id="inf43">
<mml:math id="m46">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Value added by the accommodation and catering industry</td>
<td align="center">
<inline-formula id="inf44">
<mml:math id="m47">
<mml:mrow>
<mml:mn>379.55</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">284.75</td>
<td align="char" char=".">1880.50</td>
<td align="char" char=".">13.10</td>
<td align="center">
<inline-formula id="inf45">
<mml:math id="m48">
<mml:mrow>
<mml:mn>339.47</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">100 million RMB</td>
</tr>
<tr>
<td align="left">
<inline-formula id="inf46">
<mml:math id="m49">
<mml:mrow>
<mml:mi>C</mml:mi>
<mml:mi>a</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">City area</td>
<td align="center">
<inline-formula id="inf47">
<mml:math id="m50">
<mml:mrow>
<mml:mn>6065.09</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">4625.75</td>
<td align="char" char=".">23206.32</td>
<td align="char" char=".">295.00</td>
<td align="center">
<inline-formula id="inf48">
<mml:math id="m51">
<mml:mrow>
<mml:mn>5135.54</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Square kilometers</td>
</tr>
<tr>
<td align="left">
<inline-formula id="inf49">
<mml:math id="m52">
<mml:mrow>
<mml:mi>U</mml:mi>
<mml:mi>p</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Urban population density</td>
<td align="center">
<inline-formula id="inf50">
<mml:math id="m53">
<mml:mrow>
<mml:mn>2788.65</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">2584.46</td>
<td align="char" char=".">5967.00</td>
<td align="char" char=".">515.00</td>
<td align="center">
<inline-formula id="inf51">
<mml:math id="m54">
<mml:mrow>
<mml:mn>1193.25</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">people/square km</td>
</tr>
<tr>
<td align="left">
<inline-formula id="inf52">
<mml:math id="m55">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">The number of urban populations.</td>
<td align="center">
<inline-formula id="inf53">
<mml:math id="m56">
<mml:mrow>
<mml:mn>601.03</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">493.97</td>
<td align="char" char=".">3347.32</td>
<td align="char" char=".">16.30</td>
<td align="center">
<inline-formula id="inf54">
<mml:math id="m57">
<mml:mrow>
<mml:mn>477.08</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">10,000 people</td>
</tr>
<tr>
<td align="left">
<inline-formula id="inf55">
<mml:math id="m58">
<mml:mrow>
<mml:mi>D</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Urban per capita disposable income</td>
<td align="center">
<inline-formula id="inf56">
<mml:math id="m59">
<mml:mrow>
<mml:mn>2393.92</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">2112.35</td>
<td align="char" char=".">8226.00</td>
<td align="char" char=".">64.89</td>
<td align="center">
<inline-formula id="inf57">
<mml:math id="m60">
<mml:mrow>
<mml:mn>1601.20</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">RMB</td>
</tr>
<tr>
<td align="left">
<inline-formula id="inf58">
<mml:math id="m61">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>g</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Total retail sales of consumer goods</td>
<td align="center">
<inline-formula id="inf59">
<mml:math id="m62">
<mml:mrow>
<mml:mn>26277.49</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="char" char=".">25027.32</td>
<td align="char" char=".">73848.51</td>
<td align="char" char=".">9746.80</td>
<td align="center">
<inline-formula id="inf60">
<mml:math id="m63">
<mml:mrow>
<mml:mn>11190.64</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">100 million RMB</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F3" position="float">
<label>FIGURE 3</label>
<caption>
<p>Histogram plots for the different inputs and output variables used to train the ML methods. <bold>(A)</bold> is <inline-formula id="inf61">
<mml:math id="m64">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>G</mml:mi>
<mml:mi>D</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(B)</bold> is <inline-formula id="inf62">
<mml:math id="m65">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>T</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(C)</bold> is <inline-formula id="inf63">
<mml:math id="m66">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>M</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(D)</bold> is <inline-formula id="inf64">
<mml:math id="m67">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>W</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(E)</bold> is <inline-formula id="inf65">
<mml:math id="m68">
<mml:mrow>
<mml:mi>C</mml:mi>
<mml:mi>a</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(F)</bold> is <inline-formula id="inf66">
<mml:math id="m69">
<mml:mrow>
<mml:mi>U</mml:mi>
<mml:mi>p</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(G)</bold> is <inline-formula id="inf67">
<mml:math id="m70">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(H)</bold> is <inline-formula id="inf68">
<mml:math id="m71">
<mml:mrow>
<mml:mi>D</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(I)</bold> is <inline-formula id="inf69">
<mml:math id="m72">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>g</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(J)</bold> is <inline-formula id="inf70">
<mml:math id="m73">
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>W</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g003.tif"/>
</fig>
</sec>
<sec id="s2-3">
<title>Machine Learning Techniques</title>
<sec id="s2-3-1">
<title>Data Preprocessing and Re-Sampling</title>
<p>The preprocessing methods adopted include linear normalization (<italic>Range</italic>) and standard deviation normalization (<italic>Scale</italic>), as shown in <xref ref-type="disp-formula" rid="e4">Eq. 4</xref> and <xref ref-type="disp-formula" rid="e5">Eq. 5</xref> respectively. For ML models (such as KNN) that need to calculate the distance between samples, different orders of magnitude between variables will greatly affect the performance of the model. We retained the original input data in this paper (<italic>Raw</italic>), and conducted two normalization strategies of <italic>Range</italic> and <italic>Scale</italic> to reduce the influence of data&#x2019;s dimensions and skewness on the predictions. Thus, the results of the three preprocessing methods would be comparable.<disp-formula id="e4">
<mml:math id="m74">
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(4)</label>
</disp-formula>
<disp-formula id="e5">
<mml:math id="m75">
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mi>x</mml:mi>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>x</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mrow>
<mml:mrow>
<mml:msup>
<mml:mi>&#x3c3;</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(5)</label>
</disp-formula>where <inline-formula id="inf71">
<mml:math id="m76">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>i</mml:mi>
<mml:mi>n</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> represents the minimum value of variables while <inline-formula id="inf72">
<mml:math id="m77">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>m</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>x</mml:mi>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> represents the maximum value.<inline-formula id="inf73">
<mml:math id="m78">
<mml:mrow>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>x</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> represents the numerical average value and <inline-formula id="inf74">
<mml:math id="m79">
<mml:mrow>
<mml:msup>
<mml:mi>&#x3c3;</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> is the variance of each variable.</p>
<p>To minimize the deviation caused by sampling and prevent the model from over-fitting, we adopted the 10-folds cross validation method of resampling technique to create a random sample subset of input data as a training set. The remaining data was used as test set to obtain the generalization ability of the algorithms.</p>
</sec>
<sec id="s2-3-2">
<title>Metrics of the Model</title>
<p>To evaluate the performance of each machine learning algorithm, we use three metrics of the <italic>MAE</italic>, <italic>RMSE</italic> and the&#x20;coefficient of determination (<inline-formula id="inf75">
<mml:math id="m80">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>) (<xref ref-type="bibr" rid="B7">Chai et&#x20;al., 2021</xref>; <xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021</xref>). These measurements are formulated as <xref ref-type="disp-formula" rid="e6">Eqs 6</xref>&#x2013;<xref ref-type="disp-formula" rid="e8">8</xref>.<disp-formula id="e6">
<mml:math id="m81">
<mml:mrow>
<mml:mi>M</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>E</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(6)</label>
</disp-formula>
<disp-formula id="e7">
<mml:math id="m82">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>M</mml:mi>
<mml:mi>S</mml:mi>
<mml:mi>E</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:msqrt>
<mml:mrow>
<mml:munderover>
<mml:mstyle displaystyle="true">
<mml:mo>&#x2211;</mml:mo>
</mml:mstyle>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:munderover>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:mfrac>
</mml:mrow>
</mml:mrow>
</mml:msqrt>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(7)</label>
</disp-formula>
<disp-formula id="e8">
<mml:math id="m83">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
<mml:mo>&#x2212;</mml:mo>
<mml:mfrac>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
<mml:mrow>
<mml:mstyle displaystyle="true">
<mml:msubsup>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mi>n</mml:mi>
</mml:msubsup>
<mml:mrow>
<mml:msup>
<mml:mrow>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mover accent="true">
<mml:mi>x</mml:mi>
<mml:mo>&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mstyle>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(8)</label>
</disp-formula>where <inline-formula id="inf76">
<mml:math id="m84">
<mml:mi>n</mml:mi>
</mml:math>
</inline-formula> is the number of samples, <inline-formula id="inf77">
<mml:math id="m85">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the predicted response by the model, <inline-formula id="inf78">
<mml:math id="m86">
<mml:mrow>
<mml:msub>
<mml:mi>y</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> is the actual value of the response, <inline-formula id="inf79">
<mml:math id="m87">
<mml:mrow>
<mml:mover accent="true">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo stretchy="true">&#xaf;</mml:mo>
</mml:mover>
</mml:mrow>
</mml:math>
</inline-formula> is average estimated&#x20;value.</p>
</sec>
</sec>
<sec id="s2-4">
<title>Model Interpretation</title>
<p>Model interpretability is a major challenge to applications of ML methods, which has not been given enough attention in the field of ML and MSW forecasting research. To improve the interpretations of machine learning models, this paper employed the SHAP method that assigned each input variable a value reflecting its importance to predictor (<xref ref-type="bibr" rid="B28">Lundberg and Lee, 2017</xref>).</p>
<p>For socio-economic factor subset <inline-formula id="inf80">
<mml:math id="m88">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x2286;</mml:mo>
<mml:mi>F</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> (where F stands for the set of all factors), two models are trained to extract the effect of factor <italic>i</italic>. The first model <inline-formula id="inf81">
<mml:math id="m89">
<mml:mrow>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is trained with factor <italic>i</italic> while the other one <inline-formula id="inf82">
<mml:math id="m90">
<mml:mrow>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is trained without it, where <inline-formula id="inf83">
<mml:math id="m91">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula id="inf84">
<mml:math id="m92">
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> are the values of input features/socio-economic factors. Then <inline-formula id="inf85">
<mml:math id="m93">
<mml:mrow>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:math>
</inline-formula> is computed for each possible subset <inline-formula id="inf86">
<mml:math id="m94">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x2286;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mo>\</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
<mml:mo>.</mml:mo>
</mml:mrow>
</mml:math>
</inline-formula> The Shapley value of a risk factor <italic>i</italic> is calculated using <xref ref-type="disp-formula" rid="e9">Eq. 9</xref>.<disp-formula id="e9">
<mml:math id="m95">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mi>i</mml:mi>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x2286;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mo>\</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:munder>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
<mml:mo>!</mml:mo>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>1</mml:mn>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>!</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
<mml:mo>!</mml:mo>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mstyle>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(9)</label>
</disp-formula>
</p>
<p>However, a major limitation of <xref ref-type="disp-formula" rid="e9">Eq. 9</xref> is that as the number of features/socio-economic factors increases, the computation cost will grow exponentially. To solve this problem (<xref ref-type="bibr" rid="B27">Lundberg et&#x20;al., 2020</xref>), proposed a computation-tractable explanation method, i.e.,&#x20;TreeExplainer, for decision tree-based ML models such as RF. The TreeExplainer method marks it much more efficient to calculate a risk factor&#x2019;s SHAP value both locally and globally (<xref ref-type="bibr" rid="B3">Ayoub et&#x20;al., 2021</xref>).</p>
<p>The SHAP combines optimal allocation with local explanations using the classic Shapley values. It would help users to trust the predictive models, not only what the prediction is but also why and how the prediction is made (<xref ref-type="bibr" rid="B3">Ayoub et&#x20;al., 2021</xref>). Thus, the SHAP interaction values can be calculated as the difference between the Shapley values of factor <italic>i</italic> with and without factor <italic>j</italic> in <xref ref-type="disp-formula" rid="e10">Eq. 10</xref>:<disp-formula id="e10">
<mml:math id="m96">
<mml:mrow>
<mml:msub>
<mml:mi>&#x3d5;</mml:mi>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
</mml:mrow>
</mml:msub>
<mml:mo>&#x3d;</mml:mo>
<mml:mstyle displaystyle="true">
<mml:munder>
<mml:mo>&#x2211;</mml:mo>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x2286;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mo>\</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:munder>
<mml:mrow>
<mml:mfrac>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
<mml:mo>!</mml:mo>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:mn>2</mml:mn>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>!</mml:mo>
</mml:mrow>
<mml:mrow>
<mml:mrow>
<mml:mo>&#x7c;</mml:mo>
<mml:mi>F</mml:mi>
<mml:mo>&#x7c;</mml:mo>
</mml:mrow>
<mml:mo>!</mml:mo>
</mml:mrow>
</mml:mfrac>
</mml:mrow>
</mml:mstyle>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mrow>
<mml:mi>i</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>j</mml:mi>
</mml:mrow>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>i</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>j</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mo>&#x222a;</mml:mo>
<mml:mrow>
<mml:mo>{</mml:mo>
<mml:mi>j</mml:mi>
<mml:mo>}</mml:mo>
</mml:mrow>
</mml:mrow>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>&#x2212;</mml:mo>
<mml:msub>
<mml:mi>f</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
<mml:mrow>
<mml:mo>(</mml:mo>
<mml:mrow>
<mml:msub>
<mml:mi>x</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
</mml:mrow>
<mml:mo>)</mml:mo>
</mml:mrow>
<mml:mo>.</mml:mo>
</mml:mrow>
</mml:math>
<label>(10)</label>
</disp-formula>
</p>
<p>For this superiority, we employ it to explain RF models which is based on decision trees. Therefore, compared with the existing methods (<xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021</xref>), SHAP can reflect the influence of features in each sample, show the positive and negative effects of the influence, and thereby improve the explanatory of the model output.</p>
</sec>
</sec>
<sec sec-type="results" id="s3">
<title>Results</title>
<sec id="s3-1">
<title>Comparison of Model Results</title>
<p>The programming environment used in this study is Python (version 3.8.3) with additional support packages namely scikit-learn (version 0.24.1), Tensorflow (version 2.2.2) to calculate and run the ML algorithms.</p>
<sec id="s3-1-1">
<title>Tuning</title>
<p>In this section, parameters of machine learning models are tuned, excluding multiple linear regression approach because it doesn&#x2019;t involve any hyper-parameters. Specific adjustment for parameters is shown in <xref ref-type="table" rid="T2">Table&#x20;2</xref>.</p>
<table-wrap id="T2" position="float">
<label>TABLE 2</label>
<caption>
<p>Hyper-parameters optimization.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Algorithm</th>
<th align="center">Hyper-parameters</th>
<th align="center">Other parameter settings</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">SVR</td>
<td align="left">(<inline-formula id="inf87">
<mml:math id="m97">
<mml:mrow>
<mml:mi>C</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>&#x3b5;</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>)</td>
<td align="left">Kernel &#x3d; Gaussian Kernel</td>
</tr>
<tr>
<td align="left">KNN</td>
<td align="left">K</td>
<td align="left">Using Default Parameters</td>
</tr>
<tr>
<td align="left">RF</td>
<td align="left">(<inline-formula id="inf88">
<mml:math id="m98">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>N</mml:mi>
<mml:mi>f</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>)</td>
<td align="left">Using Default Parameters</td>
</tr>
<tr>
<td align="left">XGBoost</td>
<td align="left">
<inline-formula id="inf89">
<mml:math id="m99">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</td>
<td align="left">Learning Rate &#x3d; 0.05</td>
</tr>
<tr>
<td align="left">DNN</td>
<td align="left">(<inline-formula id="inf90">
<mml:math id="m100">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>h</mml:mi>
<mml:mn>1</mml:mn>
<mml:mo>,</mml:mo>
<mml:mi>N</mml:mi>
<mml:mi>h</mml:mi>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>)</td>
<td align="left">Activation Function &#x3d; Relu</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>In the tuning process of SVR, we conduct the aforementioned three data preprocessing strategies (the Raw, Range, and Scale) respectively. As shown in <xref ref-type="table" rid="T3">Table&#x20;3</xref>, in the Raw strategy, that is to retain the original form of input data, the penalty parameter (<inline-formula id="inf91">
<mml:math id="m101">
<mml:mi>C</mml:mi>
</mml:math>
</inline-formula>) varies from 1 to 4000, compared with that in the Range strategy of 0.01&#x2013;10. The precision parameter (<italic>&#x3b5;</italic>) is an interval between 0.0001 and 0.001 in the Range and Scale strategies, compared with that of an interval from 0 to 5000. The kernel width (<inline-formula id="inf92">
<mml:math id="m102">
<mml:mi>&#x3b3;</mml:mi>
</mml:math>
</inline-formula>) doesn&#x2019;t show any differences among the three strategies. The processing strategies of Range and Scale can effectively improve the normalization and scaling of the distributions of input variables.where <inline-formula id="inf93">
<mml:math id="m103">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula id="inf94">
<mml:math id="m104">
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> in <inline-formula id="inf95">
<mml:math id="m105">
<mml:mi>&#x3b3;</mml:mi>
</mml:math>
</inline-formula> represent the results of <xref ref-type="disp-formula" rid="e11">Eq. 11</xref> and <xref ref-type="disp-formula" rid="e12">Eq. 12</xref> as the <inline-formula id="inf96">
<mml:math id="m106">
<mml:mi>&#x3b3;</mml:mi>
</mml:math>
</inline-formula> value of the SVR.<disp-formula id="e11">
<mml:math id="m107">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
<mml:mo>:</mml:mo>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mrow>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
<mml:mo>&#xd7;</mml:mo>
<mml:msup>
<mml:mi>S</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:mfrac>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(11)</label>
</disp-formula>
<disp-formula id="e12">
<mml:math id="m108">
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
<mml:mo>:</mml:mo>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mfrac>
<mml:mn>1</mml:mn>
<mml:mrow>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
</mml:mrow>
</mml:mfrac>
<mml:mo>&#xa0;</mml:mo>
<mml:mo>,</mml:mo>
</mml:mrow>
</mml:math>
<label>(12)</label>
</disp-formula>where <inline-formula id="inf97">
<mml:math id="m109">
<mml:mrow>
<mml:msub>
<mml:mi>N</mml:mi>
<mml:mi>S</mml:mi>
</mml:msub>
</mml:mrow>
</mml:math>
</inline-formula> represents the number of sample features and <inline-formula id="inf98">
<mml:math id="m110">
<mml:mrow>
<mml:msup>
<mml:mi>S</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> represents sample variance. The optimization results are shown in <xref ref-type="fig" rid="F4">Figure&#x20;4</xref>.</p>
<table-wrap id="T3" position="float">
<label>TABLE 3</label>
<caption>
<p>Hyper-parameters search space of SVR.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Strategy</th>
<th align="center">
<inline-formula id="inf99">
<mml:math id="m111">
<mml:mi>C</mml:mi>
</mml:math>
</inline-formula>
</th>
<th align="center">
<inline-formula id="inf100">
<mml:math id="m112">
<mml:mi>&#x3b3;</mml:mi>
</mml:math>
</inline-formula>
</th>
<th align="center">
<inline-formula id="inf101">
<mml:math id="m113">
<mml:mi>&#x3b5;</mml:mi>
</mml:math>
</inline-formula>
</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">Raw</td>
<td align="center">(1, 4000)</td>
<td align="left">(<inline-formula id="inf102">
<mml:math id="m114">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>)</td>
<td align="center">(0, 5000)</td>
</tr>
<tr>
<td align="left">Range</td>
<td align="center">(0.01, 10)</td>
<td align="left">(<inline-formula id="inf103">
<mml:math id="m115">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>)</td>
<td align="center">(0.0001, 0.001)</td>
</tr>
<tr>
<td align="left">Scale</td>
<td align="center">(0.01, 10)</td>
<td align="left">(<inline-formula id="inf104">
<mml:math id="m116">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
<mml:mo>,</mml:mo>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>)</td>
<td align="center">(0.0001, 0.001)</td>
</tr>
</tbody>
</table>
</table-wrap>
<fig id="F4" position="float">
<label>FIGURE 4</label>
<caption>
<p>Grid search results of SVR under different preprocess methods and different <inline-formula id="inf105">
<mml:math id="m117">
<mml:mi>&#x3b3;</mml:mi>
</mml:math>
</inline-formula>. <bold>(A)</bold> is <inline-formula id="inf106">
<mml:math id="m118">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>w</mml:mi>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mo>&#x26;</mml:mo>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(B)</bold> is <inline-formula id="inf107">
<mml:math id="m119">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mo>&#x26;</mml:mo>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(C)</bold> is <inline-formula id="inf108">
<mml:math id="m120">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mo>&#x26;</mml:mo>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(E)</bold> is <inline-formula id="inf109">
<mml:math id="m121">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>w</mml:mi>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mo>&#x26;</mml:mo>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(F)</bold> is <inline-formula id="inf110">
<mml:math id="m122">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mo>&#x26;</mml:mo>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(G)</bold> is <inline-formula id="inf111">
<mml:math id="m123">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>e</mml:mi>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mo>&#x26;</mml:mo>
<mml:mtext>&#xa0;</mml:mtext>
<mml:mi>&#x3b3;</mml:mi>
<mml:mo>&#x3d;</mml:mo>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g004.tif"/>
</fig>
<p>The hyper-parameters in other ML models are also tuned.&#x20;For RF, the number of variables tried in each segmentation (<inline-formula id="inf112">
<mml:math id="m124">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>f</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>s</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>) is set as positive integers between 1 and 9 in terms of nine input variables in this paper. The forest size (<inline-formula id="inf113">
<mml:math id="m125">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>) is set as positive integers between (50,400). The optimization results of hyper-parameters are shown in <xref ref-type="fig" rid="F5">Figure&#x20;5</xref>. In <xref ref-type="fig" rid="F4">Figures 4</xref>, <xref ref-type="fig" rid="F5">5</xref>, the redder the color is, the higher the <inline-formula id="inf114">
<mml:math id="m126">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> of the parameter combination (therefore, the better the prediction), and vice versa. For KNN, the number of neighbors <inline-formula id="inf115">
<mml:math id="m127">
<mml:mi>K</mml:mi>
</mml:math>
</inline-formula> is set as a positive integer between 1 and 29. For the XGBoost, the number of trees (<inline-formula id="inf116">
<mml:math id="m128">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>r</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>) is set to 23 positive integers between 50 and 490. For DNN, the number of neurons in the first hidden layer (<inline-formula id="inf117">
<mml:math id="m129">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>h</mml:mi>
<mml:mn>1</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>) is set as a positive integer increasing by 16 between (16,240), and the number of neurons in the second hidden layer (<inline-formula id="inf118">
<mml:math id="m130">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>h</mml:mi>
<mml:mn>2</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>) is set as one half of the number of the first hidden&#x20;layer.</p>
<fig id="F5" position="float">
<label>FIGURE 5</label>
<caption>
<p>Grid search results of RF under different preprocess methods. <bold>(A)</bold> is <inline-formula id="inf119">
<mml:math id="m131">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>w</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(B)</bold> is <inline-formula id="inf120">
<mml:math id="m132">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(C)</bold> is <inline-formula id="inf121">
<mml:math id="m133">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g005.tif"/>
</fig>
<p>Moreover, the Adma method is used as the optimization method, MAE is set as the loss function and the maximum&#x20;number of epochs is set to 200. Meanwhile, to prevent over-fitting of the DNN, the EarlyStop mechanism is introduced, and the minimum learning rate is set as 0.003 and the tolerance is set as 20. The hyper-parameter selection results of KNN, XGBoost, and DNN are shown in <xref ref-type="fig" rid="F6">Figure&#x20;6</xref>. The hyper-parameters adopted by each method are shown in&#x20;<xref ref-type="table" rid="T4">Table&#x20;4</xref>.</p>
<fig id="F6" position="float">
<label>FIGURE 6</label>
<caption>
<p>Hyperparameter optimization results of different methods under different preprocess approaches. <bold>(A)</bold> is KNN and <inline-formula id="inf122">
<mml:math id="m134">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>w</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(B)</bold> is KNN and <inline-formula id="inf123">
<mml:math id="m135">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(C)</bold> is KNN and <inline-formula id="inf124">
<mml:math id="m136">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(D)</bold> is XGBoost and <inline-formula id="inf125">
<mml:math id="m137">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>w</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(E)</bold> is XGBoost and <inline-formula id="inf126">
<mml:math id="m138">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(F)</bold> is XGBoost and <inline-formula id="inf127">
<mml:math id="m139">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(G)</bold> is DNN and <inline-formula id="inf128">
<mml:math id="m140">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>w</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(H)</bold> is DNN and <inline-formula id="inf129">
<mml:math id="m141">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(I)</bold> is DNN and <inline-formula id="inf130">
<mml:math id="m142">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g006.tif"/>
</fig>
<table-wrap id="T4" position="float">
<label>TABLE 4</label>
<caption>
<p>Hyper-parameter selection result for each algorithm.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Algorithm</th>
<th align="center">Raw</th>
<th align="center">Scale</th>
<th align="center">Range</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">SVR</td>
<td align="center">(4000,<inline-formula id="inf131">
<mml:math id="m143">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
<mml:mi>d</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>,202)</td>
<td align="center">(1.019,<inline-formula id="inf132">
<mml:math id="m144">
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, 0.0001)</td>
<td align="center">(4.049,<inline-formula id="inf133">
<mml:math id="m145">
<mml:mrow>
<mml:mi>A</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>t</mml:mi>
<mml:mi>o</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, 0.0001)</td>
</tr>
<tr>
<td align="left">KNN</td>
<td align="center">7</td>
<td align="center">3</td>
<td align="center">6</td>
</tr>
<tr>
<td align="left">RF</td>
<td align="center">(92,2)</td>
<td align="center">(78,1)</td>
<td align="center">(67,1)</td>
</tr>
<tr>
<td align="left">XGBoost</td>
<td align="center">110</td>
<td align="center">110</td>
<td align="center">170</td>
</tr>
<tr>
<td align="left">DNN</td>
<td align="center">(208,104)</td>
<td align="center">(208,104)</td>
<td align="center">(48,24)</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3-1-2">
<title>Model Application and Generation Ability</title>
<p>
<xref ref-type="fig" rid="F7">Figure&#x20;7</xref> presents the prediction performance of different ML models by using three preprocessing strategies. Several findings can conclude from the comparison among models. First, the prediction performance of MLR is the worst among all the methods because it doesn&#x2019;t involve hyper-parameter and responding adjustments. Second, the overall performances of SVR and KNN are similar, but the prediction ability of SVR is slightly higher than that of KNN except for results in Scale processing. Normally, the conducting SVR model needs a more complex process than KNN. By inputting different forms of data, the KNN only needs to adjust one super parameter, which requires less work than SVR. Third, the RF and XGBoost models present significant and similar advantages in predicting MSW production compare with MLR, SVR, and KNN according to the performance measurement of <inline-formula id="inf134">
<mml:math id="m146">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>. Fourth, the DNN has the best predictive performance among all the algorithms.</p>
<fig id="F7" position="float">
<label>FIGURE 7</label>
<caption>
<p>Comparisons of algorithms predicts performance under different preprocess methods. <bold>(A)</bold> is <inline-formula id="inf135">
<mml:math id="m147">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>w</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(B)</bold> is <inline-formula id="inf136">
<mml:math id="m148">
<mml:mrow>
<mml:mi>S</mml:mi>
<mml:mi>c</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>l</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, <bold>(C)</bold> is <inline-formula id="inf137">
<mml:math id="m149">
<mml:mrow>
<mml:mi>R</mml:mi>
<mml:mi>a</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>g</mml:mi>
<mml:mi>e</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>.</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g007.tif"/>
</fig>
<p>In this study, the RF and DNN models showed high <inline-formula id="inf138">
<mml:math id="m150">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> values (<inline-formula id="inf139">
<mml:math id="m151">
<mml:mrow>
<mml:mo>&#x3e;</mml:mo>
<mml:mn>0.9</mml:mn>
</mml:mrow>
</mml:math>
</inline-formula>) during all preprocessing methods. That means the developed ML models had a good power of explanation and were not over-fitted or over-trained. Compared with the ML method for MSW prediction developed in the earlier studies, our results were significantly better in prediction accuracy. For example (<xref ref-type="bibr" rid="B36">Niu et&#x20;al., 2021</xref>), developed LSTM and ANN models for predicting MSW generation and during the testing phase, the <inline-formula id="inf140">
<mml:math id="m152">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> value were 0.92 and 0.74, respectively (<xref ref-type="table" rid="T5">Table&#x20;5</xref>). In addition, (<xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021)</xref>, reported a DNN model with predictive performance (<inline-formula id="inf141">
<mml:math id="m153">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula>) of 0.9 for MSW production projections in Vietnam. According to <xref ref-type="bibr" rid="B21">Kumar et&#x20;al. (2018)</xref> and <xref ref-type="bibr" rid="B20">Kannangara et&#x20;al. (2018)</xref> the ANN, SVM and other ML models for predicting MSW generation showed <inline-formula id="inf142">
<mml:math id="m154">
<mml:mrow>
<mml:msup>
<mml:mi>R</mml:mi>
<mml:mn>2</mml:mn>
</mml:msup>
</mml:mrow>
</mml:math>
</inline-formula> even lower than 0.8. Thus, the machine learning model developed in this paper promotes the effective prediction of MSW production.</p>
<table-wrap id="T5" position="float">
<label>TABLE 5</label>
<caption>
<p>Comparison of model performance for prediction of MSW generation.</p>
</caption>
<table>
<thead valign="top">
<tr>
<th align="left">Method</th>
<th align="center">MAE</th>
<th align="center">RMSE</th>
<th align="center">R<sup>2</sup>
</th>
<th align="center">References</th>
</tr>
</thead>
<tbody valign="top">
<tr>
<td align="left">DNN</td>
<td align="char" char=".">861.03</td>
<td align="char" char=".">1288.80</td>
<td align="char" char=".">0.97</td>
<td rowspan="3" align="left">This study</td>
</tr>
<tr>
<td align="left">RF</td>
<td align="char" char=".">774.30</td>
<td align="char" char=".">1348.63</td>
<td align="char" char=".">0.91</td>
</tr>
<tr>
<td align="left">XGBoost</td>
<td align="char" char=".">1219.91</td>
<td align="char" char=".">1706.78</td>
<td align="char" char=".">0.90</td>
</tr>
<tr>
<td align="left">LSTM</td>
<td align="center">N/A</td>
<td align="char" char=".">935.08</td>
<td align="char" char=".">0.92</td>
<td rowspan="2" align="left">
<xref ref-type="bibr" rid="B36">Niu et&#x20;al. (2021)</xref>
</td>
</tr>
<tr>
<td align="left">ANN</td>
<td align="center">N/A</td>
<td align="char" char=".">547.14</td>
<td align="char" char=".">0.74</td>
</tr>
<tr>
<td align="left">DNN</td>
<td align="char" char=".">177.6</td>
<td align="char" char=".">294.6</td>
<td align="char" char=".">0.91</td>
<td align="left">
<xref ref-type="bibr" rid="B35">Nguyen et&#x20;al. (2021)</xref>
</td>
</tr>
<tr>
<td align="left">ANN</td>
<td align="center">N/A</td>
<td align="char" char=".">9.53</td>
<td align="char" char=".">0.75</td>
<td rowspan="3" align="left">
<xref ref-type="bibr" rid="B21">Kumar et&#x20;al. (2018)</xref>
</td>
</tr>
<tr>
<td align="left">SVM</td>
<td align="center">N/A</td>
<td align="char" char=".">9.88</td>
<td align="char" char=".">0.74</td>
</tr>
<tr>
<td align="left">RF</td>
<td align="center">N/A</td>
<td align="char" char=".">9.88</td>
<td align="char" char=".">0.66</td>
</tr>
<tr>
<td align="left">Decision Trees</td>
<td align="center">N/A</td>
<td align="char" char=".">23</td>
<td align="char" char=".">0.54</td>
<td rowspan="2" align="left">
<xref ref-type="bibr" rid="B20">Kannangara et&#x20;al. (2018)</xref>
</td>
</tr>
<tr>
<td align="left">Neural Networks</td>
<td align="center">N/A</td>
<td align="char" char=".">16</td>
<td align="char" char=".">0.72</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
</sec>
<sec id="s3-2">
<title>SHAP Analysis</title>
<sec id="s3-2-1">
<title>Overall Analysis</title>
<p>
<xref ref-type="fig" rid="F8">Figure&#x20;8</xref> shows the SHAP summary plot that orders features based on their importance to predict MSW production. Specifically, a higher SHAP value of a feature indicates higher-ranked importance to the MSW production volume. For example, the difference in the region&#x2019;s GDP has the greatest impact on the model&#x2019;s prediction of MSW production. It is likely because waste production is highly related to the household wealth that directly influences one&#x2019;s daily consumption and potential production of MSW (<xref ref-type="bibr" rid="B29">Malinauskaite et&#x20;al., 2017</xref>). Moreover, higher value of this feature result in higher SHAP values, which correspond to a higher output amount of&#x20;MSW.</p>
<fig id="F8" position="float">
<label>FIGURE 8</label>
<caption>
<p>SHAP summary&#x20;plot.</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g008.tif"/>
</fig>
<p>In addition, the industry structure presents a great influence on MSW production because of its indirect impacts on the citizens&#x2019; consumption. For instance, a higher degree of the added value of wholesale and retail trade indicates higher production of MSW compared with other industries (e.g., transportation, warehousing, and postal services industries). Some studies have argued that consumption patterns and population increase are important factors that contribute to MSW production in developing countries (<xref ref-type="bibr" rid="B25">Liu et&#x20;al., 2019</xref>; <xref ref-type="bibr" rid="B35">Nguyen et&#x20;al., 2021</xref>). Besides, the urban population also shows a significant impact on MSW production, because of its functioning on the total amount of MSW production. In contrast, other socio-economic features have a relatively insignificant impact on MSW in China. In the following paper, we will continue to analyze the dependency among these three features to discover the generation mode of MSW in China.</p>
</sec>
<sec id="s3-2-2">
<title>Dependence Analysis</title>
<p>
<xref ref-type="fig" rid="F9">Figure&#x20;9</xref> plots the relationship between a feature and its SHAP value dependent on another feature in the RF model. We select <inline-formula id="inf143">
<mml:math id="m155">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula id="inf144">
<mml:math id="m156">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>W</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> as the features to discuss and identify their variation as changes of <inline-formula id="inf145">
<mml:math id="m157">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>G</mml:mi>
<mml:mi>D</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>. As shown in <xref ref-type="fig" rid="F9">Figures 9A,B</xref>, the red points represent a higher value of <italic>InGDP</italic>, and the blue points represent the lower&#x20;one.</p>
<fig id="F9" position="float">
<label>FIGURE 9</label>
<caption>
<p>Feature dependence analysis. <bold>(A)</bold> is <inline-formula id="inf146">
<mml:math id="m158">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula id="inf147">
<mml:math id="m159">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>G</mml:mi>
<mml:mi>D</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>. <bold>(B)</bold> is <inline-formula id="inf148">
<mml:math id="m160">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>W</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> and <inline-formula id="inf149">
<mml:math id="m161">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>G</mml:mi>
<mml:mi>D</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>
</p>
</caption>
<graphic xlink:href="fenrg-09-763977-g009.tif"/>
</fig>
<p>
<xref ref-type="fig" rid="F9">Figure&#x20;9A</xref> plots the moderating effects of GDP on the impacts of urban population on MSW production. It shows that under the condition of a low <inline-formula id="inf150">
<mml:math id="m162">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> and a low <inline-formula id="inf151">
<mml:math id="m163">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>G</mml:mi>
<mml:mi>D</mml:mi>
<mml:mi>P</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula>, the SHAP value of <inline-formula id="inf152">
<mml:math id="m164">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> is below zero, which indicates that the impact of <inline-formula id="inf153">
<mml:math id="m165">
<mml:mrow>
<mml:mi>N</mml:mi>
<mml:mi>u</mml:mi>
<mml:mi>p</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> would negatively impact the MSW production under these circumstances. In other words, the less developed region might undermine the impact of the urban population on MSW production, although the local urban population increases. In contrast, with the economic growth, the increase of the urban population will promote the production of MSW. It could be recognized by the red color of the SHAP value in this figure.</p>
<p>
<xref ref-type="fig" rid="F9">Figure&#x20;9B</xref> reflects the interaction between GDP and the added value of wholesale and retail industries on MSW production. For example, before <inline-formula id="inf154">
<mml:math id="m166">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>W</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> reached 600 billion, its SHAP value is always negative. However, if <inline-formula id="inf155">
<mml:math id="m167">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>W</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> exceeds 600 billion yuan as the increase of total GDP, the increase of the added value of wholesale and retail trade plays a positive role in promoting the production of MSW. It means that if the added value of the wholesale and retail industry remains at a low level (less than 6,000 billion yuan), these industries have little effect on MSW production. However, if the added value is more than the threshold of 6000 billion yuan, the regional GDP would promote the impact of the WAR industry added value. Correspondingly, the SHAP value of <inline-formula id="inf156">
<mml:math id="m168">
<mml:mrow>
<mml:mi>I</mml:mi>
<mml:mi>n</mml:mi>
<mml:mi>W</mml:mi>
<mml:mi>A</mml:mi>
<mml:mi>R</mml:mi>
</mml:mrow>
</mml:math>
</inline-formula> indicates a significant promotion on MSW production.</p>
</sec>
</sec>
</sec>
<sec sec-type="conclusion" id="s4">
<title>Conclusion</title>
<p>To address the prediction in the production of municipal solid waste and support the WtE system design, we mainly constructed the MSW prediction method in China by using machine learning algorithms. In the comparisons of six ML models, we concentrated our attention on the predictive performances of each algorithm, particularly, by introducing three preprocessing strategies. As a result, SVR had the lowest hyperparameter consistency under different preprocessing strategies. Among the six ML methods established in this study, DNN has the best predictive ability, with an R-square of over 0.97 under all three data preprocessing strategies. The prediction performance of the machine learning methods developed in this paper is also significantly higher than the current standard (MLR) in China.</p>
<p>In addition, we find that the form of input hyper-parameter had a great influence on the models&#x2019; performances. Specifically, the explanatory indicators of the regional GDP, urban population, the added values of wholesale and retail industries, are the most important variables that affect MSW production in different provinces of China. With the development of the urban economy, the urban population increase will promote the generation of municipal solid waste. Inversely, in less developed regions, the increase of the urban population will reduce the generation of MSW. Besides, the different stages of the development of the wholesale and retail industries also impact the production of MSW. It means that in the less developed regions, a less added value of the wholesale and retail industries indicates a weak impact on MSW production, and vice&#x20;versa.</p>
<p>Our findings provide a reliable forecasting method for stakeholders. By increasing the prediction capability of MSW production, national and local policymakers could effectively conduct a series of governance policies to promote a friendly residential environment and urban sustainability. However, if given data from lower administrative, we can build even more powerful predictive models. Future studies can make effort on this to achieve more reliable and accurate results.</p>
</sec>
</body>
<back>
<sec id="s5">
<title>Data Availability Statement</title>
<p>The original contributions presented in the study are included in the article/Supplementary Material, further inquiries can be directed to the corresponding author.</p>
</sec>
<sec id="s6">
<title>Author Contributions</title>
<p>JW and LY conceived, designed, and performed the experiments. YZ, XN, and ZS prepared, analyzed the data. QG contributes policy suggestions. LY and XN wrote the early version of the paper and all authors contributed discussion and revisions, all authors have read and approved the final manuscript.</p>
</sec>
<sec id="s7">
<title>Funding</title>
<p>This research is supported by Beijing Social Science Foundation (No. 17GLB014), National Key Research and Development Program of China (2018YFF0214804), BUCT Funds for First-Class Discipline Construction (XK1802-5), BUCT (G-JD202002).</p>
</sec>
<sec sec-type="COI-statement" id="s8">
<title>Conflict of Interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec sec-type="disclaimer" id="s9">
<title>Publisher&#x2019;s Note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Abbasi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>El Hanandeh</surname>
<given-names>A.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Forecasting Municipal Solid Waste Generation Using Artificial Intelligence Modelling Approaches</article-title>. <source>Waste Manag.</source> <volume>56</volume>, <fpage>13</fpage>&#x2013;<lpage>22</lpage>. <pub-id pub-id-type="doi">10.1016/j.wasman.2016.05.018</pub-id> </citation>
</ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Al-Dahidi</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Ayadi</surname>
<given-names>O.</given-names>
</name>
<name>
<surname>Adeeb</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Louzazni</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Assessment of Artificial Neural Networks Learning Algorithms and Training Datasets for Solar Photovoltaic Power Production Prediction</article-title>. <source>Front. Energ. Res.</source> <volume>7</volume>, <fpage>130</fpage>. <pub-id pub-id-type="doi">10.3389/fenrg.2019.00130</pub-id> </citation>
</ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ayoub</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>X. J.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Combat COVID-19 Infodemic Using Explainable Natural Language Processing Models</article-title>. <source>Inf. Process. Manag.</source> <volume>58</volume>, <fpage>102569</fpage>. <pub-id pub-id-type="doi">10.1016/j.ipm.2021.102569</pub-id> </citation>
</ref>
<ref id="B4">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Beigl</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Lebersorger</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Salhofer</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2008</year>). <article-title>Modelling Municipal Solid Waste Generation: A Review</article-title>. <source>Waste Manag.</source> <volume>28</volume>, <fpage>200</fpage>&#x2013;<lpage>214</lpage>. <pub-id pub-id-type="doi">10.1016/j.wasman.2006.12.011</pub-id> </citation>
</ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Birgen</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Magnanelli</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Carlsson</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Skreiberg</surname>
<given-names>&#xd8;.</given-names>
</name>
<name>
<surname>Mosby</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Becidan</surname>
<given-names>M.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Machine Learning Based Modelling for Lower Heating Value Prediction of Municipal Solid Waste</article-title>. <source>Fuel</source> <volume>283</volume>, <fpage>118906</fpage>. <pub-id pub-id-type="doi">10.1016/j.fuel.2020.118906</pub-id> </citation>
</ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Breiman</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>1996</year>). <article-title>Bagging Predictors</article-title>. <source>Mach Learn.</source> <volume>24</volume>, <fpage>123</fpage>&#x2013;<lpage>140</lpage>. <pub-id pub-id-type="doi">10.1007/BF00058655</pub-id> </citation>
</ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Chai</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Hu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Z. G.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Structural Analysis and Forecast of Gold price Returns</article-title>. <source>J.&#x20;Manag. Sci. Eng.</source> <volume>6</volume>, <fpage>135</fpage>&#x2013;<lpage>145</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmse.2021.02.011</pub-id> </citation>
</ref>
<ref id="B8">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Chen</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Guestrin</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2016</year>). <article-title>Xgboost: A Scalable Tree Boosting System</article-title>. in <conf-name>Proceedings of the 22nd Acm Sigkdd International Conference on Knowledge Discovery and Data Mining</conf-name>, <conf-date>August 13-17, 2016</conf-date>. <publisher-loc>San Francisco, CA, USA</publisher-loc>, <fpage>785</fpage>&#x2013;<lpage>794</lpage>. </citation>
</ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Costa</surname>
<given-names>A. B. R.</given-names>
</name>
<name>
<surname>Ferreira</surname>
<given-names>P. C. G.</given-names>
</name>
<name>
<surname>Gaglianone</surname>
<given-names>W. P.</given-names>
</name>
<name>
<surname>Guill&#xe9;n</surname>
<given-names>O. T. C.</given-names>
</name>
<name>
<surname>Issler</surname>
<given-names>J.&#x20;V.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>Y.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Machine Learning and Oil price point and Density Forecasting</article-title>. <source>Energ. Econ.</source> <volume>102</volume>, <fpage>105494</fpage>. <pub-id pub-id-type="doi">10.1016/j.eneco.2021.105494</pub-id> </citation>
</ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Cover</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Hart</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>1967</year>). <article-title>Nearest Neighbor Pattern Classification</article-title>. <source>IEEE Trans. Inform. Theor.</source> <volume>13</volume>, <fpage>21</fpage>&#x2013;<lpage>27</lpage>. <pub-id pub-id-type="doi">10.1109/TIT.1967.1053964</pub-id> </citation>
</ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Demir</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Bruzzone</surname>
<given-names>L.</given-names>
</name>
</person-group> (<year>2014</year>). <article-title>A Multiple Criteria Active Learning Method for Support Vector Regression</article-title>. <source>Pattern Recognition</source> <volume>47</volume>, <fpage>2558</fpage>&#x2013;<lpage>2567</lpage>. <pub-id pub-id-type="doi">10.1016/j.patcog.2014.02.001</pub-id> </citation>
</ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ding</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>J.-W.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Cheng</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Zhao</surname>
<given-names>J.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>A Review of China&#x27;s Municipal Solid Waste (MSW) and Comparison with International Regions: Management and Technologies in Treatment and Resource Utilization</article-title>. <source>J.&#x20;Clean. Prod.</source> <volume>293</volume>, <fpage>126144</fpage>. <pub-id pub-id-type="doi">10.1016/j.jclepro.2021.126144</pub-id> </citation>
</ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Feng</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Duan</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Yakkali</surname>
<given-names>S. S.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Space Cooling Energy Usage Prediction Based on Utility Data for Residential Buildings Using Machine Learning Methods</article-title>. <source>Appl. Energ.</source> <volume>291</volume>, <fpage>116814</fpage>. <pub-id pub-id-type="doi">10.1016/j.apenergy.2021.116814</pub-id> </citation>
</ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Guo</surname>
<given-names>H.-n.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>S.-b.</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>Y.-j.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>H.-t.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Application of Machine Learning Methods for the Prediction of Organic Solid Waste Treatment and Recycling Processes: A Review</article-title>. <source>Bioresour. Tech.</source> <volume>319</volume>, <fpage>124114</fpage>. <pub-id pub-id-type="doi">10.1016/j.biortech.2020.124114</pub-id> </citation>
</ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hariharan</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Random forest Regression Analysis on Combined Role of Meteorological Indicators in Disease Dissemination in an Indian City: A Case Study of New Delhi</article-title>. <source>Urban Clim.</source> <volume>36</volume>, <fpage>100780</fpage>. <pub-id pub-id-type="doi">10.1016/j.uclim.2021.100780</pub-id> </citation>
</ref>
<ref id="B16">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>He</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Lin</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Assessment of Waste Incineration Power with Considerations of Subsidies and Emissions in China</article-title>. <source>Energy Policy</source> <volume>126</volume>, <fpage>190</fpage>&#x2013;<lpage>199</lpage>. <pub-id pub-id-type="doi">10.1016/j.enpol.2018.11.025</pub-id> </citation>
</ref>
<ref id="B17">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Hoornweg</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Bhada-Tata</surname>
<given-names>P.</given-names>
</name>
</person-group> (<year>2012</year>). <source>What a Waste: A Global Review of Solid Waste Management</source>. <comment>Urban development series; knowledge papers no. 15</comment>. <publisher-loc>Washington, DC</publisher-loc>: <publisher-name>World Bank</publisher-name>. </citation>
</ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Sun</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>A New Two-Stage Approach with Boosting and Model Averaging for Interval-Valued Crude Oil Prices Forecasting in Uncertainty Environments</article-title>. <source>Front. Energ. Res.</source> <volume>9</volume>, <fpage>707937</fpage>. <pub-id pub-id-type="doi">10.3389/fenrg.2021.707937</pub-id> </citation>
</ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Huang</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Pang</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>D.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Data-driven-based Forecasting of Two-phase Flow Parameters in Rectangular Channel</article-title>. <source>Front. Energ. Res.</source> <volume>9</volume>, <fpage>10</fpage>. <pub-id pub-id-type="doi">10.3389/fenrg.2021.641661</pub-id> </citation>
</ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kannangara</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Dua</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Ahmadi</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Bensebaa</surname>
<given-names>F.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Modeling and Prediction of Regional Municipal Solid Waste Generation and Diversion in Canada Using Machine Learning Approaches</article-title>. <source>Waste Manag.</source> <volume>74</volume>, <fpage>3</fpage>&#x2013;<lpage>15</lpage>. <pub-id pub-id-type="doi">10.1016/j.wasman.2017.11.057</pub-id> </citation>
</ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kumar</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Samadder</surname>
<given-names>S. R.</given-names>
</name>
<name>
<surname>Kumar</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Singh</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Estimation of the Generation Rate of Different Types of Plastic Wastes and Possible Revenue Recovery from Informal Recycling</article-title>. <source>Waste Manag.</source> <volume>79</volume>, <fpage>781</fpage>&#x2013;<lpage>790</lpage>. <pub-id pub-id-type="doi">10.1016/j.wasman.2018.08.045</pub-id> </citation>
</ref>
<ref id="B22">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kuznetsova</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Cardin</surname>
<given-names>M.-A.</given-names>
</name>
<name>
<surname>Diao</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>S.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Integrated Decision-Support Methodology for Combined Centralized-Decentralized Waste-To-Energy Management Systems Design</article-title>. <source>Renew. Sust. Energ. Rev.</source> <volume>103</volume>, <fpage>477</fpage>&#x2013;<lpage>500</lpage>. <pub-id pub-id-type="doi">10.1016/j.rser.2018.12.020</pub-id> </citation>
</ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname>
<given-names>R.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Tian</surname>
<given-names>W.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>On-Line Estimation Method of Lithium-Ion Battery Health Status Based on PSO-SVM</article-title>. <source>Front. Energ. Res.</source> <volume>9</volume>, <fpage>693249401</fpage>. <pub-id pub-id-type="doi">10.3389/fenrg.2021.693249</pub-id> </citation>
</ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liang</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Chai</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Y.-J.</given-names>
</name>
<name>
<surname>Zhang</surname>
<given-names>Z. G.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Refined Analysis and Prediction of Natural Gas Consumption in China</article-title>. <source>J.&#x20;Manag. Sci. Eng.</source> <volume>4</volume>, <fpage>91</fpage>&#x2013;<lpage>104</lpage>. <pub-id pub-id-type="doi">10.1016/j.jmse.2019.07.001</pub-id> </citation>
</ref>
<ref id="B25">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Li</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Gu</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>C.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>The Impact of Consumption Patterns on the Generation of Municipal Solid Waste in China: Evidences from Provincial Data</article-title>. <source>Ijerph</source> <volume>16</volume>, <fpage>1717</fpage>. <pub-id pub-id-type="doi">10.3390/ijerph16101717</pub-id> </citation>
</ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname>
<given-names>W.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>Z.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Zeng</surname>
<given-names>N.</given-names>
</name>
<name>
<surname>Liu</surname>
<given-names>Y.</given-names>
</name>
<name>
<surname>Alsaadi</surname>
<given-names>F. E.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>A Survey of Deep Neural Network Architectures and Their Applications</article-title>. <source>Neurocomputing</source> <volume>234</volume>, <fpage>11</fpage>&#x2013;<lpage>26</lpage>. <pub-id pub-id-type="doi">10.1016/j.neucom.2016.12.038</pub-id> </citation>
</ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lundberg</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Erion</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Chen</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>DeGrave</surname>
<given-names>A.</given-names>
</name>
<name>
<surname>Prutkin</surname>
<given-names>J.&#x20;M.</given-names>
</name>
<name>
<surname>Nair</surname>
<given-names>B.</given-names>
</name>
<etal/>
</person-group>&#x20;(<year>2020</year>). <article-title>From Local Explanations to Global Understanding with Explainable AI for Trees</article-title>. <source>Nat. Mach Intell.</source> <volume>2</volume>, <fpage>56</fpage>&#x2013;<lpage>67</lpage>. <pub-id pub-id-type="doi">10.1038/s42256-019-0138-9</pub-id> </citation>
</ref>
<ref id="B28">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Lundberg</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Lee</surname>
<given-names>S. I.</given-names>
</name>
</person-group> (<year>2017</year>). <article-title>A Unified Approach to Interpreting Model Predictions</article-title>. In <conf-name>31st conference on neural information processing systems</conf-name>, <fpage>4768</fpage>&#x2013;<lpage>4777</lpage>. </citation>
</ref>
<ref id="B29">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Malinauskaite</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Jouhara</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Czajczy&#x144;ska</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Stanchev</surname>
<given-names>P.</given-names>
</name>
<name>
<surname>Katsou</surname>
<given-names>E.</given-names>
</name>
<name>
<surname>Rostkowski</surname>
<given-names>P.</given-names>
</name>
<etal/>
</person-group> (<year>2017</year>). <article-title>Municipal Solid Waste Management and Waste-To-Energy in the Context of a Circular Economy and Energy Recycling in Europe</article-title>. <source>Energy</source> <volume>141</volume>, <fpage>2013</fpage>&#x2013;<lpage>2044</lpage>. <pub-id pub-id-type="doi">10.1016/j.energy.2017.11.128</pub-id> </citation>
</ref>
<ref id="B30">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mehrdad</surname>
<given-names>S. M.</given-names>
</name>
<name>
<surname>Abbasi</surname>
<given-names>M.</given-names>
</name>
<name>
<surname>Yeganeh</surname>
<given-names>B.</given-names>
</name>
<name>
<surname>Kamalan</surname>
<given-names>H.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Prediction of Methane Emission from Landfills Using Machine Learning Models</article-title>. <source>Environ. Prog. Sust. Energ.</source> <volume>40</volume>, <fpage>e13629</fpage>. <pub-id pub-id-type="doi">10.1002/ep.13629</pub-id> </citation>
</ref>
<ref id="B31">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Mukherjee</surname>
<given-names>C.</given-names>
</name>
<name>
<surname>Denney</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Mbonimpa</surname>
<given-names>E. G.</given-names>
</name>
<name>
<surname>Slagley</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Bhowmik</surname>
<given-names>R.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>A Review on Municipal Solid Waste-To-Energy Trends in the USA</article-title>. <source>Renew. Sust. Energ. Rev.</source> <volume>119</volume>, <fpage>109512</fpage>. <pub-id pub-id-type="doi">10.1016/j.rser.2019.109512</pub-id> </citation>
</ref>
<ref id="B32">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Namlis</surname>
<given-names>K.-G.</given-names>
</name>
<name>
<surname>Komilis</surname>
<given-names>D.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>Influence of Four Socioeconomic Indices and the Impact of Economic Crisis on Solid Waste Generation in Europe</article-title>. <source>Waste Manag.</source> <volume>89</volume>, <fpage>190</fpage>&#x2013;<lpage>200</lpage>. <pub-id pub-id-type="doi">10.1016/j.wasman.2019.04.012</pub-id> </citation>
</ref>
<ref id="B33">
<citation citation-type="book">
<collab>NBSC</collab> (<year>2020</year>). <source>China Statistical Yearbook 2020</source>. <publisher-loc>Beijing, China</publisher-loc>: <publisher-name>Transport and Disposal of Consumption Wastes in Cities by Region</publisher-name>. <comment>(in Chinese)</comment>. </citation>
</ref>
<ref id="B34">
<citation citation-type="book">
<collab>NBSC</collab> (<year>2021</year>). <source>Urban and Rural Population and Floating Population</source>. <publisher-loc>Beijing, China</publisher-loc>: <publisher-name>Bulletin of the Seventh National Census</publisher-name>. <comment>(No. 7) (in Chinese)</comment>. </citation>
</ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nguyen</surname>
<given-names>X. C.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>T. T. H.</given-names>
</name>
<name>
<surname>La</surname>
<given-names>D. D.</given-names>
</name>
<name>
<surname>Kumar</surname>
<given-names>G.</given-names>
</name>
<name>
<surname>Rene</surname>
<given-names>E. R.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>D. D.</given-names>
</name>
<etal/>
</person-group> (<year>2021</year>). <article-title>Development of Machine Learning - Based Models to Forecast Solid Waste Generation in Residential Areas: A Case Study from Vietnam</article-title>. <source>Resour. Conservation Recycling</source> <volume>167</volume>, <fpage>105381</fpage>. <pub-id pub-id-type="doi">10.1016/j.resconrec.2020.105381</pub-id> </citation>
</ref>
<ref id="B36">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Niu</surname>
<given-names>D.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>F.</given-names>
</name>
<name>
<surname>Dai</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>He</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Wu</surname>
<given-names>B.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Detection of Long-Term Effect in Forecasting Municipal Solid Waste Using a Long Short-Term Memory Neural Network</article-title>. <source>J.&#x20;Clean. Prod.</source> <volume>290</volume>, <fpage>125187</fpage>. <pub-id pub-id-type="doi">10.1016/j.jclepro.2020.125187</pub-id> </citation>
</ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ouda</surname>
<given-names>O. K. M.</given-names>
</name>
<name>
<surname>Cekirge</surname>
<given-names>H. M.</given-names>
</name>
<name>
<surname>Raza</surname>
<given-names>S. A. R.</given-names>
</name>
</person-group> (<year>2013</year>). <article-title>An Assessment of the Potential Contribution from Waste-To-Energy Facilities to Electricity Demand in Saudi Arabia</article-title>. <source>Energ. Convers. Manag.</source> <volume>75</volume>, <fpage>402</fpage>&#x2013;<lpage>406</lpage>. <pub-id pub-id-type="doi">10.1016/j.enconman.2013.06.056</pub-id> </citation>
</ref>
<ref id="B38">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Wu</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Kumar</surname>
<given-names>V.</given-names>
</name>
<name>
<surname>Ross Quinlan</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Ghosh</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Yang</surname>
<given-names>Q.</given-names>
</name>
<name>
<surname>Motoda</surname>
<given-names>H.</given-names>
</name>
<etal/>
</person-group> (<year>2008</year>). <article-title>Top 10 Algorithms in Data Mining</article-title>. <source>Knowl Inf. Syst.</source> <volume>14</volume>, <fpage>1</fpage>&#x2013;<lpage>37</lpage>. <pub-id pub-id-type="doi">10.1007/s10115-007-0114-2</pub-id> </citation>
</ref>
<ref id="B39">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Yang</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Nguyen</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Bui</surname>
<given-names>X.-N.</given-names>
</name>
<name>
<surname>Nguyen-Thoi</surname>
<given-names>T.</given-names>
</name>
<name>
<surname>Zhou</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Huang</surname>
<given-names>J.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Prediction of Gas Yield Generated by Energy Recovery from Municipal Solid Waste Using Deep Neural Network and Moth-Flame Optimization Algorithm</article-title>. <source>J.&#x20;Clean. Prod.</source> <volume>311</volume>, <fpage>127672</fpage>. <pub-id pub-id-type="doi">10.1016/j.jclepro.2021.127672</pub-id> </citation>
</ref>
<ref id="B40">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zhang</surname>
<given-names>X.</given-names>
</name>
<name>
<surname>Yu</surname>
<given-names>L.</given-names>
</name>
<name>
<surname>Wang</surname>
<given-names>S.</given-names>
</name>
<name>
<surname>Lai</surname>
<given-names>K. K.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Estimating the Impact of Extreme Events on Crude Oil price: An EMD-Based Event Analysis Method</article-title>. <source>Energ. Econ.</source> <volume>31</volume>, <fpage>768</fpage>&#x2013;<lpage>778</lpage>. <pub-id pub-id-type="doi">10.1016/j.eneco.2009.04.003</pub-id> </citation>
</ref>
<ref id="B41">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zheng</surname>
<given-names>J.</given-names>
</name>
<name>
<surname>Lai</surname>
<given-names>C. S.</given-names>
</name>
<name>
<surname>Yuan</surname>
<given-names>H.</given-names>
</name>
<name>
<surname>Dong</surname>
<given-names>Z. Y.</given-names>
</name>
<name>
<surname>Meng</surname>
<given-names>K.</given-names>
</name>
<name>
<surname>Lai</surname>
<given-names>L. L.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Electricity Plan Recommender System with Electrical Instruction-Based Recovery</article-title>. <source>Energy</source> <volume>203</volume>, <fpage>117775</fpage>. <pub-id pub-id-type="doi">10.1016/j.energy.2020.117775</pub-id> </citation>
</ref>
</ref-list>
</back>
</article>