<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing with OASIS Tables v3.0 20080202//EN" "https://jats.nlm.nih.gov/nlm-dtd/publishing/3.0/journalpub-oasis3.dtd">
<article xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:oasis="http://docs.oasis-open.org/ns/oasis-exchange/table" xml:lang="en" dtd-version="3.0" article-type="research-article">
  <front>
    <journal-meta><journal-id journal-id-type="publisher">GMD</journal-id><journal-title-group>
    <journal-title>Geoscientific Model Development</journal-title>
    <abbrev-journal-title abbrev-type="publisher">GMD</abbrev-journal-title><abbrev-journal-title abbrev-type="nlm-ta">Geosci. Model Dev.</abbrev-journal-title>
  </journal-title-group><issn pub-type="epub">1991-9603</issn><publisher>
    <publisher-name>Copernicus Publications</publisher-name>
    <publisher-loc>Göttingen, Germany</publisher-loc>
  </publisher></journal-meta>
    <article-meta>
      <article-id pub-id-type="doi">10.5194/gmd-19-9035-2026</article-id><title-group><article-title>A hybrid method for winter road surface temperature prediction using improved LSTMs and stacking-based ensemble learning</article-title><alt-title>Hybrid LSTM ensemble for winter road surface temperature</alt-title>
      </title-group>
      <contrib-group>
        <contrib contrib-type="author" corresp="no" rid="aff1 aff2 aff3">
          <name><surname>Li</surname><given-names>Wanting</given-names></name>
          
        <ext-link>https://orcid.org/0009-0003-4418-0576</ext-link></contrib>
        <contrib contrib-type="author" corresp="yes" rid="aff4 aff5">
          <name><surname>Zhou</surname><given-names>Linyi</given-names></name>
          <email>zhoulinyi@cma.gov.cn</email>
        </contrib>
        <contrib contrib-type="author" corresp="yes" rid="aff1 aff2 aff3">
          <name><surname>Wu</surname><given-names>Xianghua</given-names></name>
          <email>wuxianghua@nuist.edu.cn</email>
        <ext-link>https://orcid.org/0000-0002-0196-7462</ext-link></contrib>
        <contrib contrib-type="author" corresp="no" rid="aff1 aff2 aff3">
          <name><surname>Guan</surname><given-names>Yuanhong</given-names></name>
          
        </contrib>
        <contrib contrib-type="author" corresp="no" rid="aff1">
          <name><surname>Guo</surname><given-names>Yuanhao</given-names></name>
          
        </contrib>
        <contrib contrib-type="author" corresp="no" rid="aff1">
          <name><surname>Chen</surname><given-names>Kun</given-names></name>
          
        </contrib>
        <contrib contrib-type="author" corresp="no" rid="aff1">
          <name><surname>Huang</surname><given-names>Weiqi</given-names></name>
          
        </contrib>
        <contrib contrib-type="author" corresp="no" rid="aff1">
          <name><surname>Zhao</surname><given-names>Wenqian</given-names></name>
          
        </contrib>
        <aff id="aff1"><label>1</label><institution>School of Mathematics and Statistics, Nanjing University of Information Science and Technology, Nanjing 210044, China</institution>
        </aff>
        <aff id="aff2"><label>2</label><institution>Center for Applied Mathematics of Jiangsu Province, Nanjing University of Information Science and Technology, Nanjing 210044, China</institution>
        </aff>
        <aff id="aff3"><label>3</label><institution>Jiangsu International Joint Laboratory on System Modelling and Data Analysis, Nanjing University of Information Science and Technology, Nanjing 210044, China</institution>
        </aff>
        <aff id="aff4"><label>4</label><institution>Nanjing Innovation Institute for Atmospheric Sciences, Chinese Academy of Meteorological Sciences – Jiangsu Meteorological Service, Nanjing 210041, China</institution>
        </aff>
        <aff id="aff5"><label>5</label><institution>Jiangsu Key Laboratory of Transportation Meteorology of CMA/Key Laboratory of Severe Storm Disaster Risk, Nanjing 210041, China</institution>
        </aff>
      </contrib-group>
      <author-notes><corresp id="corr1">Linyi Zhou (zhoulinyi@cma.gov.cn) and Xianghua Wu (wuxianghua@nuist.edu.cn)</corresp></author-notes><pub-date><day>24</day><month>September</month><year>2026</year></pub-date>
      
      <volume>19</volume>
      <issue>18</issue>
      <fpage>9035</fpage><lpage>9061</lpage>
      <history>
        <date date-type="received"><day>28</day><month>July</month><year>2025</year></date>
           <date date-type="rev-request"><day>20</day><month>October</month><year>2025</year></date>
           <date date-type="rev-recd"><day>1</day><month>September</month><year>2026</year></date>
           <date date-type="accepted"><day>9</day><month>September</month><year>2026</year></date>
      </history>
      <permissions>
        <copyright-statement>Copyright: © 2026 Wanting Li et al.</copyright-statement>
        <copyright-year>2026</copyright-year>
      <license license-type="open-access"><license-p>This work is licensed under the Creative Commons Attribution 4.0 International License. To view a copy of this licence, visit <ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link></license-p></license></permissions><self-uri xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026.html">This article is available from https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026.html</self-uri><self-uri xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026.pdf">The full text article is available as a PDF file from https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026.pdf</self-uri>
      <abstract><title>Abstract</title>

      <p id="d2e182">Accurate prediction of road surface temperature (RST) is essential for proactive winter road maintenance and traffic safety management. However, existing approaches – ranging from physics-based models to data-driven methods – either require detailed pavement thermal parameters that are rarely available at operational road meteorological stations, or lack the capacity to simultaneously exploit local meteorological analogues and long-range temporal dependencies in an interpretable ensemble framework. This study proposes the Improved LSTMs Ensemble with Stacking (ILES) framework, which integrates two base learners with partially complementary predictive characteristics within a stacking ensemble employing out-of-fold cross-validation. The first base learner, KNN-LSTM, augments sequential modelling with similarity-based retrieval of historically analogous meteorological states to capture locally recurrent patterns. The second, BiLSTM-MHA, combines bidirectional recurrent processing with multi-head self-attention to extract long-range temporal dependencies across a 24 h input window. Moreover, a Bayesian Ridge Regression meta-learner fuses the base-learner outputs through Evidence Maximization, yielding probabilistic forecasts with closed-form posterior predictive uncertainty at the ensemble combination layer. The framework is trained and evaluated on four consecutive winter seasons (December 2020 to February 2024) of road meteorological observations from station M9393 in the northwest inland plain of Jiangsu Province, China. Results indicate that ILES achieves lower prediction errors in general than ten other models spanning persistence forecasting, traditional machine learning, and deep learning approaches, with <inline-formula><mml:math id="M1" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup></mml:mrow></mml:math></inline-formula> reaching 0.993, 0.923, and 0.826 at 1, 3, and 6 h forecasting horizons, respectively. Among three input configurations evaluated, physics-motivated feature engineering incorporating the air–surface temperature gradient and multi-scale RST temporal tendencies outperforms both the station-only baseline and ERA5-Land reanalysis augmentation, indicating that domain-knowledge-guided feature construction provides a more effective and operationally practical input strategy at instrumented sites. SHAP-based interpretability analysis, stratified by temperature regime and diurnal cycle, confirms that the learned feature importance rankings are qualitatively consistent with the dominant drivers of RST evolution identified by surface energy balance theory. Multi-site generalization is further validated at two independent stations within the same temperate monsoon climate zone, confirming the transferability of the proposed framework across different road environments within this climate setting.</p>
  </abstract>
    
<funding-group>
<award-group id="gs1">
<funding-source>National Natural Science Foundation of China</funding-source>
<award-id>U24A20606</award-id>
<award-id>42075068</award-id>
<award-id>41975087</award-id>
<award-id>42575210</award-id>
</award-group>
<award-group id="gs2">
<funding-source>Nanjing University of Information Science and Technology</funding-source>
<award-id>2025h522</award-id>
</award-group>
</funding-group>
</article-meta>
  </front>
<body>
      

<sec id="Ch1.S1" sec-type="intro">
  <label>1</label><title>Introduction</title>
      <p id="d2e205">Accurate prediction of road surface temperature (RST) is a prerequisite for effective winter road maintenance and traffic-safety management (Shao and Lister, 1996; Nantasai and Nassiri, 2019; Zhao et al., 2020). Sub-zero pavement temperatures substantially reduce the pavement friction coefficient, increasing the risk of vehicle accidents and network-wide congestion. Because RST governs the phase state of moisture at the pavement-atmosphere interface, it serves as the primary indicator for determining whether wet road surfaces will freeze under prevailing meteorological conditions (Li et al., 2022; Nowrin and Kwon, 2022; Zhang et al., 2023).</p>
      <p id="d2e208">From a physical standpoint, RST is governed by the surface energy balance (SEB) at the pavement–atmosphere interface. Following Hermansson (2004) and Chen et al. (2019), the net energy flux at the road surface can be expressed as:

          <disp-formula id="Ch1.E1" content-type="numbered"><label>1</label><mml:math id="M2" display="block"><mml:mrow><mml:msup><mml:mi>Q</mml:mi><mml:mo>*</mml:mo></mml:msup><mml:mo>=</mml:mo><mml:msub><mml:mi>Q</mml:mi><mml:mi mathvariant="normal">h</mml:mi></mml:msub><mml:mo>+</mml:mo><mml:msub><mml:mi>Q</mml:mi><mml:mi>e</mml:mi></mml:msub><mml:mo>+</mml:mo><mml:msub><mml:mi>Q</mml:mi><mml:mi mathvariant="normal">G</mml:mi></mml:msub><mml:mo>+</mml:mo><mml:msub><mml:mi>Q</mml:mi><mml:mi mathvariant="normal">m</mml:mi></mml:msub><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

        where <inline-formula><mml:math id="M3" display="inline"><mml:mrow><mml:msup><mml:mi>Q</mml:mi><mml:mo>*</mml:mo></mml:msup></mml:mrow></mml:math></inline-formula> is the net radiation flux; <inline-formula><mml:math id="M4" display="inline"><mml:mrow><mml:msub><mml:mi>Q</mml:mi><mml:mi mathvariant="normal">h</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the turbulent sensible heat flux driven by the surface–air temperature gradient and wind speed; <inline-formula><mml:math id="M5" display="inline"><mml:mrow><mml:msub><mml:mi>Q</mml:mi><mml:mi>e</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the latent heat flux modulated by relative humidity and precipitation; <inline-formula><mml:math id="M6" display="inline"><mml:mrow><mml:msub><mml:mi>Q</mml:mi><mml:mi mathvariant="normal">G</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the pavement heat storage flux governed by the substrate's thermal conductivity and heat capacity; and <inline-formula><mml:math id="M7" display="inline"><mml:mrow><mml:msub><mml:mi>Q</mml:mi><mml:mi mathvariant="normal">m</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> represents the latent heat of fusion associated with freezing or melting of surface moisture. All flux terms adopt a sign convention in which positive values denote energy directed into the road surface and negative values denote energy loss from the surface; under this convention, <inline-formula><mml:math id="M8" display="inline"><mml:mrow><mml:msup><mml:mi>Q</mml:mi><mml:mo>*</mml:mo></mml:msup></mml:mrow></mml:math></inline-formula> represents the net radiative energy available to drive the remaining flux terms. RST reflects the cumulative outcome of these fluxes over preceding hours, exhibiting diurnal periodicity, lagged responses to radiative forcing, and abrupt thermal transitions during precipitation events or cold-air advection (Hermansson, 2004; Chen et al., 2019). This physical understanding motivates both the selection of meteorological predictors and the design of derived input features in the present study.</p>
      <p id="d2e319">RST prediction presents two fundamental challenges that existing approaches have not simultaneously resolved. Physics-based models – including finite-difference, finite-element, and finite-volume formulations – solve the surface energy balance equation coupled to one-dimensional heat conduction in the pavement substrate and can yield physically interpretable, high-accuracy solutions when site-specific thermal parameters are well characterised (Hermansson, 2004; Wang et al., 2009; Chen et al., 2019; Minhoto et al., 2005; Schindler et al., 2004; Saliko et al., 2023). However, their accuracy is critically sensitive to pavement thermal and radiative properties – thermal conductivity, volumetric heat capacity, surface emissivity, and the convective heat transfer coefficient – that are rarely available at operational road meteorological stations (Qin and Hiller, 2013; Athukorallage et al., 2023; Ayasrah et al., 2023; Adwan et al., 2021). This parameter identifiability constraint extends to operational systems such as METRo (Crevier and Delage, 2001), whose performance remains dependent on accurate specification of pavement layer properties, surface albedo, and anthropogenic heat sources that vary across road segments and degrade over time (Shao and Lister, 1996; Kangas et al., 2015; Karsisto et al., 2016), fundamentally limiting the scalability of physics-based approaches to the large number of heterogeneous road segments for which detailed subsurface characterisation is unavailable. Concurrently, RST is a strongly autocorrelated variable whose evolution reflects the integrated effect of meteorological forcing over preceding hours, exhibiting lagged responses to radiative and turbulent fluxes and regime-dependent nonlinear dynamics during precipitation events or nocturnal radiative cooling. Statistical and empirical models – including multiple linear regression, stepwise regression, and regression-based nowcast systems – have identified near-surface air temperature and lagged RST as the dominant predictors (Yin et al., 2019; Kršmanc et al., 2013; Li et al., 2018; Diefenderfer et al., 2006; Asefzadeh et al., 2017; Hassan et al., 2005), yet these approaches impose linearity or low-order polynomial constraints on the inherently nonlinear meteorology–RST relationship, precluding adequate representation of the multi-hour temporal dependencies and abrupt thermal transitions that characterise real-world RST evolution (Wang, 2015; Gedafa et al., 2014; Darghiasi et al., 2025; Jing and Zhang, 2018).</p>
      <p id="d2e322">Data-driven machine learning (ML) methods address the nonlinearity limitation by learning flexible input–output mappings directly from observational records without requiring explicit thermal parameter calibration. Ensemble tree methods, including random forest (RF), gradient boosting decision trees (GBDT), and extreme gradient boosting (XGBoost), have achieved strong predictive performance by aggregating predictions from multiple weak learners (Milad et al., 2021b; Qiu et al., 2020; Liu et al., 2018; Kebede et al., 2024; Yuan et al., 2025). Support vector regression (SVR) and Gaussian process regression (GPR) have captured nonlinear RST–meteorology relationships under varying surface conditions (Molavi Nojumi et al., 2022; Wang et al., 2023), and artificial neural networks (ANNs) have demonstrated superior pattern-learning capacity relative to conventional regression approaches (Abo-Hashema, 2013; Rigabadi et al., 2022). However, these methods treat each input sample independently and therefore cannot exploit the multi-hour temporal autocorrelation structure inherent in RST time series, constraining their ability to represent the thermal inertia and lag dynamics that are characteristic of RST evolution (Yang et al., 2020; Hatamzad et al., 2022; Darghiasi et al., 2024).</p>
      <p id="d2e326">Recurrent deep learning architectures, and in particular Long Short-Term Memory (LSTM) networks and their variants, address this limitation more directly by learning sequential dependencies through gating mechanisms that selectively retain and discard information over multi-hour input windows (Dai et al., 2023; Ghalandari et al., 2023). Several studies have demonstrated their superiority over traditional ML approaches for RST prediction: Tabrizi et al. (2021) showed that a hybrid CNN-LSTM model outperformed LSTM, ConvLSTM, Seq2Seq, and WaveNet baselines across 1, 2, 4, and 6 h horizons; Zhang et al. (2023) reported MAE of 0.82 °C and RMSE of 1.24 °C with an XGBoost-LSTNet combination; Li et al. (2022) showed that a hybrid Bayesian structural time series-Bayesian neural network model attained greater than 95 % coverage within predicted confidence intervals; Dai et al. (2023) proposed a GRU-LSTM ensemble exploiting RST periodicity and meteorological lag effects; and Zhang et al. (2024) integrated RF-based feature selection with LSTM for short-term prediction at 10 min resolution. Bidirectional LSTM (BiLSTM) architectures, which process the input sequence in both forward and backward temporal directions, have shown enhanced feature-extraction capabilities relative to unidirectional LSTMs (Maddu et al., 2021; Tao et al., 2024; Milad et al., 2021a). Bai et al. (2022) demonstrated that an attention-based Bi-LSTM achieved 93.4 % of predictions within 1 °C error. Table 1 summarizes recent LSTM-based approaches for RST prediction, including the model architectures employed, the input features considered, and representative predictive accuracy reported in each study.</p>

<table-wrap id="T1" specific-use="star"><label>Table 1</label><caption><p id="d2e332">Summary of LSTM-based approaches for road surface temperature prediction. The accuracy values reported in Table 1 are drawn from the original publications and reflect different datasets, input variables, temporal resolutions, and evaluation protocols, they are presented here for literature-positioning purposes only. Abbreviations: AT, air temperature; SR, solar radiation; RH, relative humidity; <inline-formula><mml:math id="M9" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula>, precipitation; WS, wind speed; WD, wind direction; AP, atmospheric pressure; Depth, measurement depth below pavement surface; Time, time-of-day index.</p></caption><oasis:table frame="topbot"><oasis:tgroup cols="6">
     <oasis:colspec colnum="1" colname="col1" align="justify" colwidth="2.5cm"/>
     <oasis:colspec colnum="2" colname="col2" align="justify" colwidth="2cm"/>
     <oasis:colspec colnum="3" colname="col3" align="justify" colwidth="2cm"/>
     <oasis:colspec colnum="4" colname="col4" align="justify" colwidth="1cm"/>
     <oasis:colspec colnum="5" colname="col5" align="justify" colwidth="4cm"/>
     <oasis:colspec colnum="6" colname="col6" align="justify" colwidth="3cm"/>
     <oasis:thead>
       <oasis:row rowsep="1">
         <oasis:entry colname="col1" align="left">Reference</oasis:entry>
         <oasis:entry colname="col2" align="left">Model</oasis:entry>
         <oasis:entry colname="col3" align="left">Feature</oasis:entry>
         <oasis:entry colname="col4" align="left">Interval</oasis:entry>
         <oasis:entry colname="col5" align="left">Characteristic</oasis:entry>
         <oasis:entry colname="col6" align="left">Evaluation</oasis:entry>
       </oasis:row>
     </oasis:thead>
     <oasis:tbody>
       <oasis:row rowsep="1">
         <oasis:entry colname="col1" align="left">Tabrizi et al. (2021)</oasis:entry>
         <oasis:entry colname="col2" align="left">CNN-LSTM</oasis:entry>
         <oasis:entry colname="col3" align="left">RST, AT, SR</oasis:entry>
         <oasis:entry colname="col4" align="left">1 h</oasis:entry>
         <oasis:entry colname="col5" align="left">CNN-LSTM hybrid architecture for multi-horizon forecasting</oasis:entry>
         <oasis:entry colname="col6" align="left">MAE <inline-formula><mml:math id="M10" display="inline"><mml:mo>=</mml:mo></mml:math></inline-formula> 1.05–3.43  °C and <inline-formula><mml:math id="M11" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup><mml:mo>=</mml:mo></mml:mrow></mml:math></inline-formula> 0.98–0.80 across 1 to 6 h horizons</oasis:entry>
       </oasis:row>
       <oasis:row rowsep="1">
         <oasis:entry colname="col1" align="left">Milad et al. (2021a)</oasis:entry>
         <oasis:entry colname="col2" align="left">Bi-LSTM</oasis:entry>
         <oasis:entry colname="col3" align="left">AT, Depth, Time</oasis:entry>
         <oasis:entry colname="col4" align="left">1 h</oasis:entry>
         <oasis:entry colname="col5" align="left">Bidirectional LSTM with enhanced feature extraction across depth and time dimensions</oasis:entry>
         <oasis:entry colname="col6" align="left">MAE <inline-formula><mml:math id="M12" display="inline"><mml:mo>=</mml:mo></mml:math></inline-formula> 1.332  °C; <inline-formula><mml:math id="M13" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup><mml:mo>=</mml:mo></mml:mrow></mml:math></inline-formula> 0.956</oasis:entry>
       </oasis:row>
       <oasis:row rowsep="1">
         <oasis:entry colname="col1" align="left">Bai et al. (2022)</oasis:entry>
         <oasis:entry colname="col2" align="left">Att-BiLSTM</oasis:entry>
         <oasis:entry colname="col3" align="left">RST, AT, P, WS, RH</oasis:entry>
         <oasis:entry colname="col4" align="left">1 h</oasis:entry>
         <oasis:entry colname="col5" align="left">Attention-augmented BiLSTM with sliding window optimisation for micro-scale prediction</oasis:entry>
         <oasis:entry colname="col6" align="left">MAE <inline-formula><mml:math id="M14" display="inline"><mml:mo>=</mml:mo></mml:math></inline-formula> 0.330  °C; 93.4 % of predictions within <inline-formula><mml:math id="M15" display="inline"><mml:mo>±</mml:mo></mml:math></inline-formula>1 °C</oasis:entry>
       </oasis:row>
       <oasis:row rowsep="1">
         <oasis:entry colname="col1" align="left">Dai et al. (2023)</oasis:entry>
         <oasis:entry colname="col2" align="left">GRU, LSTM</oasis:entry>
         <oasis:entry colname="col3" align="left">RST, AT, P, WS, RH</oasis:entry>
         <oasis:entry colname="col4" align="left">1 h</oasis:entry>
         <oasis:entry colname="col5" align="left">GRU-LSTM ensemble exploiting RST periodicity and meteorological lag effects</oasis:entry>
         <oasis:entry colname="col6" align="left">MAE of 0.345, 0.833, and 1.743 °C at 1, 3, and 6 h horizons</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1" align="left">Zhang et al. (2024)</oasis:entry>
         <oasis:entry colname="col2" align="left">RF-LSTM</oasis:entry>
         <oasis:entry colname="col3" align="left">RST, AT, AP, WS, RH, WD</oasis:entry>
         <oasis:entry colname="col4" align="left">10 min</oasis:entry>
         <oasis:entry colname="col5" align="left">RF-based feature selection integrated with LSTM for short-term prediction</oasis:entry>
         <oasis:entry colname="col6" align="left">MAE <inline-formula><mml:math id="M16" display="inline"><mml:mo>=</mml:mo></mml:math></inline-formula> 0.048  °C; 99.1 % within <inline-formula><mml:math id="M17" display="inline"><mml:mo>±</mml:mo></mml:math></inline-formula>0.5 °C</oasis:entry>
       </oasis:row>
     </oasis:tbody>
   </oasis:tgroup></oasis:table></table-wrap>

      <p id="d2e566">Despite these advances, the existing literature exhibits several limitations that warrant further investigation. First, individual LSTM-based architectures tend to emphasise either local pattern recurrence associated with historically similar meteorological trajectories or long-range temporal dependencies spanning the full input window, but rarely capture both simultaneously. While some studies have explored ensemble combinations of heterogeneous models, the systematic integration of architecturally complementary LSTM variants through principled meta-learning remains underexplored for RST prediction. Moreover, the consistency of ensemble-derived predictive gains across multiple forecasting horizons has not been thoroughly demonstrated. Second, the relative contribution of different input data sources to RST prediction accuracy remains insufficiently characterised. While several studies have incorporated reanalysis products or derived meteorological variables alongside station observations, systematic comparisons across input configurations under consistent experimental conditions are scarce, making it difficult to determine whether the additional data complexity is justified by commensurate predictive gains. Third, although post-hoc feature attribution methods such as SHAP (SHapley Additive exPlanations; Lundberg and Lee, 2017) have been applied to data-driven RST models, existing analyses have typically reported globally averaged attribution values without systematic stratification by temperature regime or time of day. Such stratified analysis would provide a more rigorous basis for evaluating whether learned model behaviour is qualitatively consistent with the dominant meteorological drivers identified by SEB theory, and for identifying conditions under which the model may be less reliable.</p>
      <p id="d2e569">Motivated by these limitations, this study proposes the Improved LSTMs Ensemble with Stacking (ILES) framework for winter RST prediction. Rather than relying on a single recurrent architecture, ILES integrates two base learners with partially complementary predictive characteristics – KNN-LSTM for similarity-driven local pattern retrieval and BiLSTM-MHA for long-range temporal dependency extraction – within a stacking ensemble, where a Bayesian Ridge Regression meta-learner fuses the base learner outputs to yield probabilistic forecasts with closed-form posterior predictive uncertainty at the ensemble combination layer; the ablation analysis shows that the combined gain from these two components is sub-additive. Three input configurations – a station-only baseline, ERA5-Land reanalysis augmentation, and physics-motivated feature engineering – are systematically compared to determine the most effective and operationally practical input strategy. In addition, SHAP-based attribution analysis is conducted in a stratified manner across temperature regimes and diurnal cycles, providing a more systematic evaluation of whether learned feature importance is consistent with the dominant drivers identified by SEB theory. The framework is trained and evaluated on four consecutive winter seasons of hourly road meteorological observations from station M9393, located in the northwest inland plain of Jiangsu Province, China, with multi-site generalisation assessed at two independent stations M9474 and M9448.</p>
      <p id="d2e572">The remainder of this paper is organized as follows. Section 2 describes the ILES framework, including the KNN-LSTM and BiLSTM-MHA architectures, the stacking ensemble construction, and the Bayesian Ridge Regression meta-learner. Section 3 presents the study site, observational dataset, ERA5-Land reanalysis data, feature engineering procedure, and experimental design. Section 4 reports results encompassing overall predictive performance, input configuration comparisons, ensemble complementarity analysis, and uncertainty quantification. Section 5 summarizes the principal conclusions.</p>
</sec>
<sec id="Ch1.S2">
  <label>2</label><title>Methodology</title>
      <p id="d2e583">Figure 1 presents an overview of the ILES framework for winter RST prediction, organised into three sequential stages: data preparation and feature extraction, model training and ensemble construction, and prediction performance evaluation.</p>

      <fig id="F1" specific-use="star"><label>Figure 1</label><caption><p id="d2e588">Overview of the ILES model for winter RST prediction.</p></caption>
        <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f01.png"/>

      </fig>

      <p id="d2e597">In the first stage, historical meteorological observations from station M9393 and ERA5-Land reanalysis variables are subjected to quality control procedures including outlier removal, missing value imputation, and hourly resampling. Input variables are selected on the basis of Spearman rank correlation with winter RST, and the retained variables are standardised using <inline-formula><mml:math id="M18" display="inline"><mml:mi>Z</mml:mi></mml:math></inline-formula>-score normalisation parameters estimated exclusively from the training set. A sliding window is applied to construct sequential input samples across three candidate input configurations, enabling the model to capture the multi-hour autocorrelation structure and diurnal periodicity inherent in RST time series.</p>
      <p id="d2e608">In the second stage, two structurally distinct LSTM variants serve as base learners within a stacking ensemble. KNN-LSTM augments temporal sequence modelling with similarity-based feature retrieval from historical analogues, enabling the model to leverage locally recurrent meteorological patterns. BiLSTM-MHA employs bidirectional recurrent processing combined with multi-head self-attention to extract long-range temporal dependencies across the full input window. Base learner predictions for the meta-learner training matrix are generated through three-fold temporal cross-validation aligned to complete winter seasons, with each fold corresponding to one complete winter as the validation period. A BRR meta-learner is fitted on the resulting out-of-fold predictions, yielding probabilistic forecasts with closed-form uncertainty quantification as the final ensemble output.</p>
      <p id="d2e611">In the third stage, the ILES model is assessed from four analytical perspectives: comprehensive benchmarking against ten models across 1, 3, and 6 h forecasting horizons; comparison of three input variable combinations to evaluate the relative predictive value of station observations, ERA5-Land reanalysis augmentation, and physics-motivated feature engineering; stratified SHAP-based attribution analysis across temperature regimes and diurnal cycles; and multi-site validation at stations with distinct geographical and surface conditions.</p>
<sec id="Ch1.S2.SS1">
  <label>2.1</label><title>KNN-LSTM</title>
      <p id="d2e621">The proposed KNN-LSTM model integrates similarity-based feature augmentation with temporal sequence modelling to improve predictive accuracy for wintertime RST forecasting, building upon the framework established by Luo et al. (2019) for traffic flow prediction.</p>
      <p id="d2e624">Consider a forecasting task in which the objective is to predict RST at time <inline-formula><mml:math id="M19" display="inline"><mml:mrow><mml:mi>t</mml:mi><mml:mo>+</mml:mo><mml:mi>h</mml:mi></mml:mrow></mml:math></inline-formula> from a sliding window of historical observations. Let <inline-formula><mml:math id="M20" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> denote the input feature matrix at time <inline-formula><mml:math id="M21" display="inline"><mml:mi>t</mml:mi></mml:math></inline-formula>, of dimension <inline-formula><mml:math id="M22" display="inline"><mml:mrow><mml:mi>T</mml:mi><mml:mo>×</mml:mo><mml:mi>n</mml:mi></mml:mrow></mml:math></inline-formula>, where <inline-formula><mml:math id="M23" display="inline"><mml:mi>T</mml:mi></mml:math></inline-formula> is the window size and <inline-formula><mml:math id="M24" display="inline"><mml:mi>n</mml:mi></mml:math></inline-formula> is the feature dimension. Given a training set of <inline-formula><mml:math id="M25" display="inline"><mml:mi>M</mml:mi></mml:math></inline-formula> historical samples <inline-formula><mml:math id="M26" display="inline"><mml:mrow><mml:mo mathvariant="italic">{</mml:mo><mml:mo>(</mml:mo><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo>)</mml:mo><mml:msubsup><mml:mo mathvariant="italic">}</mml:mo><mml:mrow><mml:mi>i</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>M</mml:mi></mml:msubsup></mml:mrow></mml:math></inline-formula>, where <inline-formula><mml:math id="M27" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the <inline-formula><mml:math id="M28" display="inline"><mml:mi>i</mml:mi></mml:math></inline-formula>th input window and <inline-formula><mml:math id="M29" display="inline"><mml:mrow><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the corresponding target RST, the similarity between a query window <inline-formula><mml:math id="M30" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> and each training sample is measured by Euclidean distance in the flattened feature space. The resulting distance matrix is normalised via min–max scaling to ensure numerical comparability:

            <disp-formula id="Ch1.E2" content-type="numbered"><label>2</label><mml:math id="M31" display="block"><mml:mrow><mml:msub><mml:mi mathvariant="bold">D</mml:mi><mml:mi mathvariant="normal">norm</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mrow><mml:mi mathvariant="bold">D</mml:mi><mml:mo>-</mml:mo><mml:mo>min⁡</mml:mo><mml:mfenced open="(" close=")"><mml:mi mathvariant="bold">D</mml:mi></mml:mfenced></mml:mrow><mml:mrow><mml:mo>max⁡</mml:mo><mml:mfenced close=")" open="("><mml:mi mathvariant="bold">D</mml:mi></mml:mfenced><mml:mo>-</mml:mo><mml:mo>min⁡</mml:mo><mml:mfenced close=")" open="("><mml:mi mathvariant="bold">D</mml:mi></mml:mfenced><mml:mo>+</mml:mo><mml:mi mathvariant="italic">ϵ</mml:mi></mml:mrow></mml:mfrac></mml:mstyle><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

          where <inline-formula><mml:math id="M32" display="inline"><mml:mi mathvariant="italic">ϵ</mml:mi></mml:math></inline-formula> is a small constant to avoid division by zero. The <inline-formula><mml:math id="M33" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> nearest neighbours <inline-formula><mml:math id="M34" display="inline"><mml:mrow><mml:mo mathvariant="italic">{</mml:mo><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mi>k</mml:mi></mml:mfenced></mml:mrow></mml:msub><mml:msubsup><mml:mo mathvariant="italic">}</mml:mo><mml:mrow><mml:mi>k</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>K</mml:mi></mml:msubsup></mml:mrow></mml:math></inline-formula> are identified as the <inline-formula><mml:math id="M35" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> reference samples with the smallest normalised distances to <inline-formula><mml:math id="M36" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>. The similarity feature matrix <inline-formula><mml:math id="M37" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">F</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is then constructed as their uniform average:

            <disp-formula id="Ch1.E3" content-type="numbered"><label>3</label><mml:math id="M38" display="block"><mml:mrow><mml:msub><mml:mi mathvariant="bold">F</mml:mi><mml:mi>t</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mn mathvariant="normal">1</mml:mn><mml:mi>K</mml:mi></mml:mfrac></mml:mstyle><mml:msubsup><mml:mo>∑</mml:mo><mml:mrow><mml:mi>k</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>K</mml:mi></mml:msubsup><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mi>k</mml:mi></mml:mfenced></mml:mrow></mml:msub><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

          capturing the central tendency of historically analogous meteorological states. The augmented input is formed by concatenating <inline-formula><mml:math id="M39" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> and <inline-formula><mml:math id="M40" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">F</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> along the feature dimension:

            <disp-formula id="Ch1.E4" content-type="numbered"><label>4</label><mml:math id="M41" display="block"><mml:mrow><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="normal">aug</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mo>[</mml:mo><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mi>t</mml:mi></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold">F</mml:mi><mml:mi>t</mml:mi></mml:msub><mml:mo>]</mml:mo><mml:mo>.</mml:mo></mml:mrow></mml:math></disp-formula></p>
      <p id="d2e980">This augmentation strategy preserves the original temporal structure while incorporating rich contextual information from analogous meteorological conditions in the historical record, and the dimensionality of the input doubles from <inline-formula><mml:math id="M42" display="inline"><mml:mi>n</mml:mi></mml:math></inline-formula> to 2<inline-formula><mml:math id="M43" display="inline"><mml:mi>n</mml:mi></mml:math></inline-formula>. The augmented sequence <inline-formula><mml:math id="M44" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="normal">aug</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:math></inline-formula> is fed into a three-layer LSTM network (Hochreiter and Schmidhuber, 1997) designed to capture hierarchical temporal dependencies:

                <disp-formula specific-use="gather" content-type="numbered"><mml:math id="M45" display="block"><mml:mtable displaystyle="true"><mml:mlabeledtr id="Ch1.E5"><mml:mtd><mml:mtext>5</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle class="stylechange" displaystyle="true"/><mml:msubsup><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced open="(" close=")"><mml:mn mathvariant="normal">1</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi mathvariant="bold-italic">c</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">1</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>=</mml:mo><mml:msub><mml:mi mathvariant="normal">LSTM</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub><mml:mfenced open="(" close=")"><mml:mrow><mml:msub><mml:mi mathvariant="bold">X</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="normal">aug</mml:mi></mml:mrow></mml:msub><mml:mo>,</mml:mo><mml:msubsup><mml:mi>h</mml:mi><mml:mrow><mml:mi mathvariant="italic">τ</mml:mi><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mrow><mml:mfenced open="(" close=")"><mml:mn mathvariant="normal">1</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi>c</mml:mi><mml:mrow><mml:mi mathvariant="italic">τ</mml:mi><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mrow><mml:mfenced open="(" close=")"><mml:mn mathvariant="normal">1</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>;</mml:mo><mml:msub><mml:mi mathvariant="italic">θ</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub></mml:mrow></mml:mfenced><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr><mml:mlabeledtr id="Ch1.E6"><mml:mtd><mml:mtext>6</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle class="stylechange" displaystyle="true"/><mml:msubsup><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced open="(" close=")"><mml:mn mathvariant="normal">2</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi mathvariant="bold-italic">c</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">2</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>=</mml:mo><mml:msub><mml:mi mathvariant="normal">LSTM</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msub><mml:mfenced open="(" close=")"><mml:mrow><mml:msubsup><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">1</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi>h</mml:mi><mml:mrow><mml:mi mathvariant="italic">τ</mml:mi><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mrow><mml:mfenced open="(" close=")"><mml:mn mathvariant="normal">2</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi>c</mml:mi><mml:mrow><mml:mi mathvariant="italic">τ</mml:mi><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mrow><mml:mfenced open="(" close=")"><mml:mn mathvariant="normal">2</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>;</mml:mo><mml:msub><mml:mi mathvariant="italic">θ</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msub></mml:mrow></mml:mfenced><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr><mml:mlabeledtr id="Ch1.E7"><mml:mtd><mml:mtext>7</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle class="stylechange" displaystyle="true"/><mml:msubsup><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">3</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi mathvariant="bold-italic">c</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">3</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>=</mml:mo><mml:msub><mml:mi mathvariant="normal">LSTM</mml:mi><mml:mn mathvariant="normal">3</mml:mn></mml:msub><mml:mfenced close=")" open="("><mml:mrow><mml:msubsup><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">2</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi>h</mml:mi><mml:mrow><mml:mi mathvariant="italic">τ</mml:mi><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">3</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi>c</mml:mi><mml:mrow><mml:mi mathvariant="italic">τ</mml:mi><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mrow><mml:mfenced open="(" close=")"><mml:mn mathvariant="normal">3</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>;</mml:mo><mml:msub><mml:mi mathvariant="italic">θ</mml:mi><mml:mn mathvariant="normal">3</mml:mn></mml:msub></mml:mrow></mml:mfenced><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr></mml:mtable></mml:math></disp-formula>

          where <inline-formula><mml:math id="M46" display="inline"><mml:mrow><mml:msubsup><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced open="(" close=")"><mml:mi mathvariant="normal">ℓ</mml:mi></mml:mfenced></mml:mrow></mml:msubsup></mml:mrow></mml:math></inline-formula> and <inline-formula><mml:math id="M47" display="inline"><mml:mrow><mml:msubsup><mml:mi mathvariant="bold-italic">c</mml:mi><mml:mi mathvariant="italic">τ</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mi mathvariant="normal">ℓ</mml:mi></mml:mfenced></mml:mrow></mml:msubsup></mml:mrow></mml:math></inline-formula> denote the hidden state and cell state at layer <inline-formula><mml:math id="M48" display="inline"><mml:mi mathvariant="normal">ℓ</mml:mi></mml:math></inline-formula> and time step <inline-formula><mml:math id="M49" display="inline"><mml:mi mathvariant="italic">τ</mml:mi></mml:math></inline-formula>, respectively, and <inline-formula><mml:math id="M50" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="italic">θ</mml:mi><mml:mi mathvariant="normal">ℓ</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> represents the learnable parameters. The final hidden state <inline-formula><mml:math id="M51" display="inline"><mml:mrow><mml:msubsup><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi>T</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">3</mml:mn></mml:mfenced></mml:mrow></mml:msubsup></mml:mrow></mml:math></inline-formula> is passed through a fully connected layer to generate the prediction:

            <disp-formula id="Ch1.E8" content-type="numbered"><label>8</label><mml:math id="M52" display="block"><mml:mrow><mml:msub><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo stretchy="false" mathvariant="normal">^</mml:mo></mml:mover><mml:mrow><mml:mi>t</mml:mi><mml:mo>+</mml:mo><mml:mi>h</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:msub><mml:mi mathvariant="bold">W</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mo>⋅</mml:mo><mml:msubsup><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi>T</mml:mi><mml:mrow><mml:mfenced close=")" open="("><mml:mn mathvariant="normal">3</mml:mn></mml:mfenced></mml:mrow></mml:msubsup><mml:mo>+</mml:mo><mml:msub><mml:mi>b</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

          where <inline-formula><mml:math id="M53" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">W</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub></mml:mrow></mml:math></inline-formula> and <inline-formula><mml:math id="M54" display="inline"><mml:mrow><mml:msub><mml:mi>b</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub></mml:mrow></mml:math></inline-formula> are the learnable weight matrix and bias term, respectively. A separate instance of this model, with independently optimised parameters <inline-formula><mml:math id="M55" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="italic">θ</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="italic">θ</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="italic">θ</mml:mi><mml:mn mathvariant="normal">3</mml:mn></mml:msub></mml:mrow></mml:math></inline-formula> and <inline-formula><mml:math id="M56" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">W</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>b</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub></mml:mrow></mml:math></inline-formula>, is trained for each forecasting horizon; no parameters are shared across horizons. A schematic overview of the KNN-LSTM architecture is provided in Figs. 1 and 2.</p>

      <fig id="F2"><label>Figure 2</label><caption><p id="d2e1464">Unrolled computation of a single LSTM layer across the input window, where <inline-formula><mml:math id="M57" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi mathvariant="italic">τ</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> and <inline-formula><mml:math id="M58" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold-italic">c</mml:mi><mml:mi mathvariant="italic">τ</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> denote the hidden and cell states at time step <inline-formula><mml:math id="M59" display="inline"><mml:mi mathvariant="italic">τ</mml:mi></mml:math></inline-formula>, with <inline-formula><mml:math id="M60" display="inline"><mml:mrow><mml:mi mathvariant="italic">τ</mml:mi><mml:mo>∈</mml:mo><mml:mo mathvariant="italic">{</mml:mo><mml:mn mathvariant="normal">1</mml:mn><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:mi>T</mml:mi><mml:mo mathvariant="italic">}</mml:mo></mml:mrow></mml:math></inline-formula> indexing the steps within the window of length <inline-formula><mml:math id="M61" display="inline"><mml:mi>T</mml:mi></mml:math></inline-formula>.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f02.png"/>

        </fig>

</sec>
<sec id="Ch1.S2.SS2">
  <label>2.2</label><title>BiLSTM-MHA</title>
      <p id="d2e1541">The proposed BiLSTM-MHA model integrates a Bidirectional Long Short-Term Memory network with a multi-head self-attention mechanism (Vaswani et al., 2017), augmented by residual connections and layer normalisation (Ba et al., 2016) to maintain training stability while capturing complex temporal dependencies.</p>
      <p id="d2e1544">Let the sliding window size be <inline-formula><mml:math id="M62" display="inline"><mml:mi>T</mml:mi></mml:math></inline-formula> and the input sequence at time <inline-formula><mml:math id="M63" display="inline"><mml:mi>t</mml:mi></mml:math></inline-formula> be <inline-formula><mml:math id="M64" display="inline"><mml:mrow><mml:mi mathvariant="bold">X</mml:mi><mml:mo>=</mml:mo><mml:mo>(</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">x</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>-</mml:mo><mml:mi>T</mml:mi><mml:mo>+</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">x</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>-</mml:mo><mml:mi>T</mml:mi><mml:mo>+</mml:mo><mml:mn mathvariant="normal">2</mml:mn></mml:mrow></mml:msub><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">x</mml:mi><mml:mi>t</mml:mi></mml:msub><mml:msup><mml:mo>)</mml:mo><mml:mi mathvariant="normal">⊤</mml:mi></mml:msup></mml:mrow></mml:math></inline-formula>, where <inline-formula><mml:math id="M65" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">x</mml:mi><mml:mi>t</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mo>(</mml:mo><mml:msub><mml:mi>x</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>,</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>x</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>,</mml:mo><mml:mn mathvariant="normal">2</mml:mn></mml:mrow></mml:msub><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:msub><mml:mi>x</mml:mi><mml:mrow><mml:mi>t</mml:mi><mml:mo>,</mml:mo><mml:mi>n</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:math></inline-formula>) and <inline-formula><mml:math id="M66" display="inline"><mml:mi>n</mml:mi></mml:math></inline-formula> denotes the feature dimension.</p>
      <p id="d2e1679">The BiLSTM layer captures temporal dependencies in both temporal directions (Schuster and Paliwal, 1997). At each time step <inline-formula><mml:math id="M67" display="inline"><mml:mi>t</mml:mi></mml:math></inline-formula>, the forward LSTM processes the sequence from past to future and the backward LSTM processes the same input sequence in the reverse temporal direction, producing a concatenated hidden state <inline-formula><mml:math id="M68" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi>t</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mo>[</mml:mo><mml:msub><mml:mover accent="true"><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mo mathvariant="normal">→</mml:mo></mml:mover><mml:mi>t</mml:mi></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mover accent="true"><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mo mathvariant="normal">←</mml:mo></mml:mover><mml:mi>t</mml:mi></mml:msub><mml:msup><mml:mo>]</mml:mo><mml:mi mathvariant="normal">⊤</mml:mi></mml:msup></mml:mrow></mml:math></inline-formula>, where <inline-formula><mml:math id="M69" display="inline"><mml:mrow><mml:msub><mml:mover accent="true"><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mo mathvariant="normal">→</mml:mo></mml:mover><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> and <inline-formula><mml:math id="M70" display="inline"><mml:mrow><mml:msub><mml:mover accent="true"><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mo mathvariant="normal">←</mml:mo></mml:mover><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> denote the forward and backward hidden states, respectively. The full sequence <inline-formula><mml:math id="M71" display="inline"><mml:mrow><mml:mi>H</mml:mi><mml:mo>=</mml:mo><mml:mo>(</mml:mo><mml:msub><mml:mi>h</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>h</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">h</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>) is passed to the subsequent attention layer. It should be noted that the backward pass operates exclusively within the historical input window <inline-formula><mml:math id="M72" display="inline"><mml:mrow><mml:mfenced close="]" open="["><mml:mrow><mml:mi>t</mml:mi><mml:mo>-</mml:mo><mml:mi>T</mml:mi><mml:mo>+</mml:mo><mml:mn mathvariant="normal">1</mml:mn><mml:mo>,</mml:mo><mml:mi>t</mml:mi></mml:mrow></mml:mfenced></mml:mrow></mml:math></inline-formula> and accesses no observations beyond time <inline-formula><mml:math id="M73" display="inline"><mml:mi>t</mml:mi></mml:math></inline-formula>.</p>
      <p id="d2e1820">The multi-head self-attention mechanism (Vaswani et al., 2017) transforms <inline-formula><mml:math id="M74" display="inline"><mml:mi>H</mml:mi></mml:math></inline-formula> through learned linear projections into multiple representation subspaces in parallel, enabling the model to attend simultaneously to different temporal positions and feature abstractions:

            <disp-formula id="Ch1.E9" content-type="numbered"><label>9</label><mml:math id="M75" display="block"><mml:mrow><mml:mi mathvariant="normal">MultiHead</mml:mi><mml:mfenced open="(" close=")"><mml:mrow><mml:mi mathvariant="bold">Q</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="bold">K</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="bold">V</mml:mi></mml:mrow></mml:mfenced><mml:mo>=</mml:mo><mml:mi mathvariant="normal">Concat</mml:mi><mml:mfenced close=")" open="("><mml:mrow><mml:msub><mml:mi mathvariant="normal">head</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:mi mathvariant="normal">⋯</mml:mi><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="normal">head</mml:mi><mml:mi mathvariant="normal">h</mml:mi></mml:msub></mml:mrow></mml:mfenced><mml:msup><mml:mi mathvariant="bold">W</mml:mi><mml:mi mathvariant="normal">O</mml:mi></mml:msup><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

          where each head is defined as head<inline-formula><mml:math id="M76" display="inline"><mml:mrow><mml:msub><mml:mi/><mml:mi>i</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mi mathvariant="normal">Attention</mml:mi><mml:mo>(</mml:mo><mml:mi>Q</mml:mi><mml:msubsup><mml:mi mathvariant="bold">W</mml:mi><mml:mi>i</mml:mi><mml:mi>Q</mml:mi></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi mathvariant="bold">KW</mml:mi><mml:mi>i</mml:mi><mml:mi>K</mml:mi></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi mathvariant="bold">VW</mml:mi><mml:mi>i</mml:mi><mml:mi>V</mml:mi></mml:msubsup></mml:mrow></mml:math></inline-formula>) with Attention<inline-formula><mml:math id="M77" display="inline"><mml:mrow><mml:mo>(</mml:mo><mml:mi mathvariant="bold">Q</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="bold">K</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="bold">V</mml:mi><mml:mo>)</mml:mo><mml:mo>=</mml:mo><mml:mi mathvariant="normal">softmax</mml:mi><mml:mfenced open="(" close=")"><mml:mstyle displaystyle="false"><mml:mfrac style="text"><mml:mrow><mml:msup><mml:mi mathvariant="bold">QK</mml:mi><mml:mi mathvariant="normal">⊤</mml:mi></mml:msup></mml:mrow><mml:msqrt><mml:mrow><mml:msub><mml:mi>d</mml:mi><mml:mi mathvariant="normal">k</mml:mi></mml:msub></mml:mrow></mml:msqrt></mml:mfrac></mml:mstyle></mml:mfenced><mml:mi mathvariant="bold">V</mml:mi></mml:mrow></mml:math></inline-formula>. Here <inline-formula><mml:math id="M78" display="inline"><mml:mrow><mml:msup><mml:mi mathvariant="bold">W</mml:mi><mml:mi mathvariant="normal">O</mml:mi></mml:msup></mml:mrow></mml:math></inline-formula> denotes the output projection matrix, <inline-formula><mml:math id="M79" display="inline"><mml:mrow><mml:msubsup><mml:mi mathvariant="bold">W</mml:mi><mml:mi>i</mml:mi><mml:mi>Q</mml:mi></mml:msubsup></mml:mrow></mml:math></inline-formula>, <inline-formula><mml:math id="M80" display="inline"><mml:mrow><mml:msubsup><mml:mi mathvariant="bold">W</mml:mi><mml:mi>i</mml:mi><mml:mi>K</mml:mi></mml:msubsup></mml:mrow></mml:math></inline-formula>, <inline-formula><mml:math id="M81" display="inline"><mml:mrow><mml:msubsup><mml:mi mathvariant="bold">W</mml:mi><mml:mi>i</mml:mi><mml:mi>V</mml:mi></mml:msubsup></mml:mrow></mml:math></inline-formula> are the per-head projection matrices, and <inline-formula><mml:math id="M82" display="inline"><mml:mrow><mml:msub><mml:mi>d</mml:mi><mml:mi mathvariant="normal">k</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the key dimension. In the self-attention formulation adopted here, <inline-formula><mml:math id="M83" display="inline"><mml:mi mathvariant="bold">Q</mml:mi></mml:math></inline-formula>, <inline-formula><mml:math id="M84" display="inline"><mml:mi mathvariant="bold">K</mml:mi></mml:math></inline-formula>, and <inline-formula><mml:math id="M85" display="inline"><mml:mi mathvariant="bold">V</mml:mi></mml:math></inline-formula> are all derived from <inline-formula><mml:math id="M86" display="inline"><mml:mi>H</mml:mi></mml:math></inline-formula>.</p>
      <p id="d2e2056">A residual connection is applied to the attention output to facilitate gradient flow during backpropagation (He et al., 2016), followed by layer normalisation to stabilise the feature distribution:

                <disp-formula specific-use="gather" content-type="numbered"><mml:math id="M87" display="block"><mml:mtable displaystyle="true"><mml:mlabeledtr id="Ch1.E10"><mml:mtd><mml:mtext>10</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle class="stylechange" displaystyle="true"/><mml:mi mathvariant="bold">R</mml:mi><mml:mo>=</mml:mo><mml:mi mathvariant="bold">H</mml:mi><mml:mo>+</mml:mo><mml:mi mathvariant="normal">MultiHead</mml:mi><mml:mo>(</mml:mo><mml:mi mathvariant="bold">H</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="bold">H</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="bold">H</mml:mi><mml:mo>)</mml:mo><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr><mml:mlabeledtr id="Ch1.E11"><mml:mtd><mml:mtext>11</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle displaystyle="true" class="stylechange"/><mml:mi mathvariant="bold">Z</mml:mi><mml:mo>=</mml:mo><mml:mi mathvariant="normal">LayerNorm</mml:mi><mml:mo>(</mml:mo><mml:mi mathvariant="bold">R</mml:mi><mml:mo>)</mml:mo><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr></mml:mtable></mml:math></disp-formula>

          Temporal aggregation is then performed by global average pooling across all time steps:

            <disp-formula id="Ch1.E12" content-type="numbered"><label>12</label><mml:math id="M88" display="block"><mml:mrow><mml:mi mathvariant="bold-italic">c</mml:mi><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mn mathvariant="normal">1</mml:mn><mml:mi>T</mml:mi></mml:mfrac></mml:mstyle><mml:msubsup><mml:mo>∑</mml:mo><mml:mrow><mml:mi>t</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>T</mml:mi></mml:msubsup><mml:msub><mml:mi mathvariant="bold">Z</mml:mi><mml:mi>t</mml:mi></mml:msub><mml:mo>.</mml:mo></mml:mrow></mml:math></disp-formula></p>
      <p id="d2e2154">Global average pooling reduces parameter count relative to concatenation-based aggregation and confers robustness to sequence length variation (Lin et al., 2013). The final prediction is generated through a fully connected output layer:

            <disp-formula id="Ch1.E13" content-type="numbered"><label>13</label><mml:math id="M89" display="block"><mml:mrow><mml:msub><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo mathvariant="normal" stretchy="false">^</mml:mo></mml:mover><mml:mrow><mml:mi>t</mml:mi><mml:mo>+</mml:mo><mml:mi>h</mml:mi></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:msub><mml:mi mathvariant="bold">W</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mi>c</mml:mi><mml:mo>+</mml:mo><mml:msub><mml:mi>b</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

          where <inline-formula><mml:math id="M90" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">W</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub></mml:mrow></mml:math></inline-formula> represents the weight matrix and <inline-formula><mml:math id="M91" display="inline"><mml:mrow><mml:msub><mml:mi>b</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub></mml:mrow></mml:math></inline-formula> denotes the bias term. As with KNN-LSTM, a separate instance of the BiLSTM-MHA model – including all BiLSTM, attention, and output-layer parameters – is trained independently for each forecasting horizon, with no weight sharing across horizons. A schematic of the BiLSTM-MHA architecture is provided in Figs. 1 and 3.</p>

      <fig id="F3"><label>Figure 3</label><caption><p id="d2e2218">Multi-Head Attention consists of <inline-formula><mml:math id="M92" display="inline"><mml:mi>h</mml:mi></mml:math></inline-formula> parallel attention heads, each operating on a distinct learned projection of the input. <inline-formula><mml:math id="M93" display="inline"><mml:mi mathvariant="bold">Q</mml:mi></mml:math></inline-formula>, <inline-formula><mml:math id="M94" display="inline"><mml:mi mathvariant="bold">K</mml:mi></mml:math></inline-formula>, and <inline-formula><mml:math id="M95" display="inline"><mml:mi mathvariant="bold">V</mml:mi></mml:math></inline-formula> are all derived from the BiLSTM output.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f03.png"/>

        </fig>

</sec>
<sec id="Ch1.S2.SS3">
  <label>2.3</label><title>Stacking ensemble with out-of-fold cross-validation</title>
      <p id="d2e2263">Stacking is a hierarchical ensemble learning method that combines predictions from multiple base learners through a meta-learner to achieve superior predictive performance (Wolpert, 1992; Breiman, 1996). Unlike bagging and boosting, which rely on parallel or sequential resampling strategies, stacking employs a two-layer architecture where diverse base models generate intermediate predictions that are subsequently integrated by a meta-model. This framework has demonstrated remarkable success in time series forecasting tasks by leveraging model complementarity (Divina et al., 2018).</p>
      <p id="d2e2266">A fundamental requirement of stacking ensemble methods is that the meta-learner must be trained on predictions that the base learners generated for samples they did not observe during their own training. Violating this requirement introduces data leakage: base learners that are first trained on the full training set and subsequently used to predict in-sample observations produce fitted values rather than genuine forecasts, causing the meta-learner to learn a combination rule optimized for interpolation rather than generalization (Wolpert, 1992). To satisfy this requirement, we generate base learner predictions for the meta-learner's training matrix through a <inline-formula><mml:math id="M96" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula>-fold temporal cross-validation scheme. Specifically, the model consists of two layers: the first layer comprises <inline-formula><mml:math id="M97" display="inline"><mml:mi>N</mml:mi></mml:math></inline-formula> distinct base learners, while the second layer incorporates a meta-learner. The multi-model fusion prediction workflow of Stacking, based on <inline-formula><mml:math id="M98" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula>-fold cross-validation, can be summarized as follows: <list list-type="order"><list-item>
      <p id="d2e2292"><italic>Dataset Partitioning.</italic> The original dataset is first partitioned into a training set <inline-formula><mml:math id="M99" display="inline"><mml:mi>D</mml:mi></mml:math></inline-formula> and a holdout test set <inline-formula><mml:math id="M100" display="inline"><mml:mi>B</mml:mi></mml:math></inline-formula>. The training set <inline-formula><mml:math id="M101" display="inline"><mml:mi>D</mml:mi></mml:math></inline-formula> is further subdivided into <inline-formula><mml:math id="M102" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> disjoint subsets, denoted <inline-formula><mml:math id="M103" display="inline"><mml:mrow><mml:msub><mml:mi>D</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>D</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:msub><mml:mi>D</mml:mi><mml:mi>K</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>.</p></list-item><list-item>
      <p id="d2e2355"><italic>Base Learner Training.</italic> For each base learner <inline-formula><mml:math id="M104" display="inline"><mml:mrow><mml:msub><mml:mi>L</mml:mi><mml:mi>n</mml:mi></mml:msub><mml:mfenced open="(" close=")"><mml:mrow><mml:mi>n</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">2</mml:mn><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:mi>N</mml:mi></mml:mrow></mml:mfenced></mml:mrow></mml:math></inline-formula>, <inline-formula><mml:math id="M105" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula>-fold cross-validation is performed. In the <inline-formula><mml:math id="M106" display="inline"><mml:mi>k</mml:mi></mml:math></inline-formula>th fold, the validation set is <inline-formula><mml:math id="M107" display="inline"><mml:mrow><mml:msub><mml:mi>D</mml:mi><mml:mi>k</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> and the training set is <inline-formula><mml:math id="M108" display="inline"><mml:mrow><mml:msub><mml:mi>D</mml:mi><mml:mrow><mml:mo>(</mml:mo><mml:mo>-</mml:mo><mml:mi>k</mml:mi><mml:mo>)</mml:mo></mml:mrow></mml:msub><mml:mo>=</mml:mo><mml:mi>D</mml:mi><mml:mi mathvariant="normal">∖</mml:mi><mml:msub><mml:mi>D</mml:mi><mml:mi>k</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>. This procedure yields <inline-formula><mml:math id="M109" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> independently trained instances of each base learner, one per fold.</p></list-item><list-item>
      <p id="d2e2454"><italic>Construction of the Augmented Training Set.</italic> Each trained instance of base learner <inline-formula><mml:math id="M110" display="inline"><mml:mrow><mml:msub><mml:mi>L</mml:mi><mml:mi>n</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> generates out-of-fold predictions on its corresponding validation subset <inline-formula><mml:math id="M111" display="inline"><mml:mrow><mml:msub><mml:mi>D</mml:mi><mml:mi>k</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>. Concatenating these predictions across all <inline-formula><mml:math id="M112" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> folds produce a full-length out-of-fold prediction vector <inline-formula><mml:math id="M113" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold-italic">T</mml:mi><mml:mi>n</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> for base learner <inline-formula><mml:math id="M114" display="inline"><mml:mrow><mml:msub><mml:mi>L</mml:mi><mml:mi>n</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>:<disp-formula id="Ch1.E14" content-type="numbered"><label>14</label><mml:math id="M115" display="block"><mml:mrow><mml:msub><mml:mi mathvariant="bold-italic">T</mml:mi><mml:mi>n</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mo mathvariant="italic">{</mml:mo><mml:msub><mml:mi>t</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>t</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>t</mml:mi><mml:mn mathvariant="normal">3</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:msub><mml:mi>t</mml:mi><mml:mi>K</mml:mi></mml:msub><mml:mo mathvariant="italic">}</mml:mo><mml:mo>.</mml:mo></mml:mrow></mml:math></disp-formula></p>
      <p id="d2e2560">The out-of-fold prediction vectors from all <inline-formula><mml:math id="M116" display="inline"><mml:mi>N</mml:mi></mml:math></inline-formula> base learners are concatenated with the original training set <inline-formula><mml:math id="M117" display="inline"><mml:mi>D</mml:mi></mml:math></inline-formula> to form the augmented training set <inline-formula><mml:math id="M118" display="inline"><mml:mrow><mml:msup><mml:mi>D</mml:mi><mml:mo>′</mml:mo></mml:msup></mml:mrow></mml:math></inline-formula>, which serves as input for the second-level meta-learner:<disp-formula id="Ch1.E15" content-type="numbered"><label>15</label><mml:math id="M119" display="block"><mml:mrow><mml:msup><mml:mi>D</mml:mi><mml:mo>′</mml:mo></mml:msup><mml:mo>=</mml:mo><mml:mo mathvariant="italic">{</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">T</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">T</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">T</mml:mi><mml:mn mathvariant="normal">3</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">T</mml:mi><mml:mi>N</mml:mi></mml:msub><mml:mo>,</mml:mo><mml:mi>D</mml:mi><mml:mo mathvariant="italic">}</mml:mo><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula></p>
      <p id="d2e2641">Upon completion of <inline-formula><mml:math id="M120" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula>-fold cross-validation, each base learner <inline-formula><mml:math id="M121" display="inline"><mml:mrow><mml:msub><mml:mi>L</mml:mi><mml:mi>n</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> produces <inline-formula><mml:math id="M122" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> prediction vectors on the holdout test set <inline-formula><mml:math id="M123" display="inline"><mml:mi>B</mml:mi></mml:math></inline-formula>, denoted <inline-formula><mml:math id="M124" display="inline"><mml:mrow><mml:msubsup><mml:mi mathvariant="bold-italic">b</mml:mi><mml:mi>n</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msubsup><mml:mo>,</mml:mo><mml:msubsup><mml:mi mathvariant="bold-italic">b</mml:mi><mml:mi>n</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msubsup><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:msubsup><mml:mi mathvariant="bold-italic">b</mml:mi><mml:mi>n</mml:mi><mml:mi>K</mml:mi></mml:msubsup></mml:mrow></mml:math></inline-formula>. These <inline-formula><mml:math id="M125" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> predictions are averaged to yield a single representative test prediction <inline-formula><mml:math id="M126" display="inline"><mml:mrow><mml:msub><mml:mi>b</mml:mi><mml:mi>n</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> for base learner <inline-formula><mml:math id="M127" display="inline"><mml:mrow><mml:msub><mml:mi>L</mml:mi><mml:mi>n</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>. The consolidated test predictions from all <inline-formula><mml:math id="M128" display="inline"><mml:mi>N</mml:mi></mml:math></inline-formula> base learners are combined with <inline-formula><mml:math id="M129" display="inline"><mml:mi>B</mml:mi></mml:math></inline-formula> to form the augmented test set <inline-formula><mml:math id="M130" display="inline"><mml:mrow><mml:msup><mml:mi>B</mml:mi><mml:mo>′</mml:mo></mml:msup></mml:mrow></mml:math></inline-formula>:<disp-formula id="Ch1.E16" content-type="numbered"><label>16</label><mml:math id="M131" display="block"><mml:mrow><mml:msup><mml:mi>B</mml:mi><mml:mo>′</mml:mo></mml:msup><mml:mo>=</mml:mo><mml:mo mathvariant="italic">{</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">b</mml:mi><mml:mn mathvariant="normal">1</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">b</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">b</mml:mi><mml:mn mathvariant="normal">3</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:mi mathvariant="normal">…</mml:mi><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">b</mml:mi><mml:mi>N</mml:mi></mml:msub><mml:mo>,</mml:mo><mml:mi>B</mml:mi><mml:mo mathvariant="italic">}</mml:mo><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula></p></list-item><list-item>
      <p id="d2e2821"><italic>Meta-Learner Training.</italic> The augmented dataset <inline-formula><mml:math id="M132" display="inline"><mml:mrow><mml:msup><mml:mi>D</mml:mi><mml:mo>′</mml:mo></mml:msup></mml:mrow></mml:math></inline-formula> is used to train the second-level meta-learner, and <inline-formula><mml:math id="M133" display="inline"><mml:mrow><mml:msup><mml:mi>B</mml:mi><mml:mo>′</mml:mo></mml:msup></mml:mrow></mml:math></inline-formula> is used to evaluate its generalisation performance. The meta-learner learns an optimal combination strategy over the base learner outputs, yielding a final ensemble prediction that is expected to surpass any individual constituent model.</p></list-item></list></p>
      <p id="d2e2848">The number of folds <inline-formula><mml:math id="M134" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> is determined by the temporal structure of the data. Since the training dataset spans three complete winter seasons, a 3-fold configuration naturally implements a leave-one-season-out protocol in which each fold corresponds to exactly one complete winter as the validation period. This structure mirrors the real-world forecasting scenario of predicting an unseen winter season from prior observations, and fold boundaries are strictly aligned with seasonal transitions with temporal ordering preserved throughout. The complete workflow is illustrated in Fig. 4.</p>

      <fig id="F4" specific-use="star"><label>Figure 4</label><caption><p id="d2e2861">Architecture of the Stacking Ensemble based on Out-of-Fold Cross-Validation.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f04.png"/>

        </fig>

</sec>
<sec id="Ch1.S2.SS4">
  <label>2.4</label><title>Bayesian ridge regression meta-learner</title>
      <p id="d2e2878">Since the meta-learner operates on base learner predictions rather than raw input features, its primary role is to learn an optimal fusion of these predictions that accounts for their relative strengths and inter-model correlations. We adopt Bayesian Ridge Regression (BRR; MacKay, 1992) as the meta-learner, motivated by three considerations. First, base learner predictions tend to be correlated due to shared input features; the L2 regularization inherent in BRR effectively alleviates multicollinearity (Hoerl and Kennard, 1970). Second, the Bayesian framework automatically determines the optimal regularization strength through Evidence Maximization, precluding overfitting without manual hyperparameter tuning (Bishop, 2006). Third, BRR yields probabilistic outputs that quantify prediction uncertainty at the ensemble combination layer, conditional on the deterministic base-learner outputs (Gelman and Shalizi, 2013).</p>
      <p id="d2e2881">The BRR meta-learner assumes a linear relationship between base learner predictions and the observed RST, with Gaussian noise and hierarchical priors on the weight precision:

                <disp-formula specific-use="gather" content-type="numbered"><mml:math id="M135" display="block"><mml:mtable displaystyle="true"><mml:mlabeledtr id="Ch1.E17"><mml:mtd><mml:mtext>17</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle displaystyle="true" class="stylechange"/><mml:mi mathvariant="bold-italic">y</mml:mi><mml:mo>=</mml:mo><mml:mi mathvariant="bold">X</mml:mi><mml:mi mathvariant="bold-italic">w</mml:mi><mml:mo>+</mml:mo><mml:mi mathvariant="bold-italic">ϵ</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="bold-italic">ϵ</mml:mi><mml:mo>∼</mml:mo><mml:mi mathvariant="script">N</mml:mi><mml:mo>(</mml:mo><mml:mn mathvariant="normal">0</mml:mn><mml:mo>,</mml:mo><mml:msup><mml:mi mathvariant="italic">α</mml:mi><mml:mrow><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow></mml:msup><mml:mi mathvariant="bold">I</mml:mi><mml:mo>)</mml:mo><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr><mml:mlabeledtr id="Ch1.E18"><mml:mtd><mml:mtext>18</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle displaystyle="true" class="stylechange"/><mml:mtable rowspacing="0.2ex" class="split" displaystyle="true" columnalign="right left"><mml:mtr><mml:mtd/><mml:mtd><mml:mrow><mml:mi mathvariant="bold-italic">w</mml:mi><mml:mo>∼</mml:mo><mml:mi mathvariant="script">N</mml:mi><mml:mo>(</mml:mo><mml:mn mathvariant="normal">0</mml:mn><mml:mo>,</mml:mo><mml:msup><mml:mi mathvariant="italic">λ</mml:mi><mml:mrow><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow></mml:msup><mml:mi mathvariant="bold">I</mml:mi><mml:mo>)</mml:mo><mml:mo>,</mml:mo><mml:mi mathvariant="italic">α</mml:mi><mml:mo>∼</mml:mo><mml:mi mathvariant="normal">Gamma</mml:mi><mml:mo>(</mml:mo><mml:msub><mml:mi>a</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>b</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mo>)</mml:mo><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr><mml:mtr><mml:mtd/><mml:mtd><mml:mrow><mml:mi mathvariant="italic">λ</mml:mi><mml:mo>∼</mml:mo><mml:mi mathvariant="normal">Gamma</mml:mi><mml:mo>(</mml:mo><mml:msub><mml:mi>c</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi>d</mml:mi><mml:mn mathvariant="normal">0</mml:mn></mml:msub><mml:mo>)</mml:mo><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr></mml:mtable></mml:mrow></mml:mtd></mml:mlabeledtr></mml:mtable></mml:math></disp-formula></p>
      <p id="d2e3028">The posterior distribution over weights is analytically tractable:

            <disp-formula id="Ch1.E19" content-type="numbered"><label>19</label><mml:math id="M136" display="block"><mml:mtable class="split" rowspacing="0.2ex" displaystyle="true" columnalign="right left"><mml:mtr><mml:mtd/><mml:mtd><mml:mrow><mml:mi mathvariant="bold-italic">w</mml:mi><mml:mo>∣</mml:mo><mml:mi mathvariant="bold-italic">y</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="bold">X</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="italic">α</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="italic">λ</mml:mi><mml:mo>∼</mml:mo><mml:mi mathvariant="script">N</mml:mi><mml:mo>(</mml:mo><mml:msub><mml:mi mathvariant="bold-italic">μ</mml:mi><mml:mi>n</mml:mi></mml:msub><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold">Σ</mml:mi><mml:mi>n</mml:mi></mml:msub><mml:mo>)</mml:mo><mml:mo>,</mml:mo><mml:msub><mml:mi mathvariant="bold">Σ</mml:mi><mml:mi>n</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mo>(</mml:mo><mml:mi mathvariant="italic">λ</mml:mi><mml:mi mathvariant="bold">I</mml:mi><mml:mo>+</mml:mo><mml:mi mathvariant="italic">α</mml:mi><mml:msup><mml:mi mathvariant="bold">X</mml:mi><mml:mi mathvariant="normal">⊤</mml:mi></mml:msup><mml:mi mathvariant="bold">X</mml:mi><mml:msup><mml:mo>)</mml:mo><mml:mrow><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow></mml:msup><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mtr><mml:mtr><mml:mtd/><mml:mtd><mml:mrow><mml:mspace linebreak="nobreak" width="0.25em"/><mml:mspace width="0.25em" linebreak="nobreak"/><mml:mspace linebreak="nobreak" width="0.25em"/><mml:mspace linebreak="nobreak" width="0.25em"/><mml:mspace linebreak="nobreak" width="0.25em"/><mml:msub><mml:mi mathvariant="bold-italic">μ</mml:mi><mml:mi>n</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mi mathvariant="italic">α</mml:mi><mml:msub><mml:mi mathvariant="bold">Σ</mml:mi><mml:mi>n</mml:mi></mml:msub><mml:msup><mml:mi mathvariant="bold">X</mml:mi><mml:mi mathvariant="normal">⊤</mml:mi></mml:msup><mml:mi mathvariant="bold-italic">y</mml:mi><mml:mo>.</mml:mo></mml:mrow></mml:mtd></mml:mtr></mml:mtable></mml:math></disp-formula></p>
      <p id="d2e3149">The precision hyperparameters <inline-formula><mml:math id="M137" display="inline"><mml:mi mathvariant="italic">α</mml:mi></mml:math></inline-formula> and <inline-formula><mml:math id="M138" display="inline"><mml:mi mathvariant="italic">λ</mml:mi></mml:math></inline-formula> are optimised via Type-II Maximum Likelihood (Evidence Maximisation), which automatically balances data fit against regularisation without requiring manual cross-validation.</p>
      <p id="d2e3167">Beyond point prediction, BRR yields a closed-form posterior predictive distribution for each forecast:

            <disp-formula id="Ch1.E20" content-type="numbered"><label>20</label><mml:math id="M139" display="block"><mml:mrow><mml:mi>p</mml:mi><mml:mo>(</mml:mo><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo stretchy="false" mathvariant="normal">̃</mml:mo></mml:mover><mml:mo>∣</mml:mo><mml:mi>x</mml:mi><mml:mo>,</mml:mo><mml:mi mathvariant="script">D</mml:mi><mml:mo>)</mml:mo><mml:mo>=</mml:mo><mml:mi mathvariant="script">N</mml:mi><mml:mo>(</mml:mo><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo mathvariant="normal" stretchy="false">̃</mml:mo></mml:mover><mml:mo>∣</mml:mo><mml:msubsup><mml:mi mathvariant="italic">μ</mml:mi><mml:mi>n</mml:mi><mml:mi mathvariant="normal">⊤</mml:mi></mml:msubsup><mml:mi mathvariant="bold-italic">x</mml:mi><mml:mo>,</mml:mo><mml:msubsup><mml:mi mathvariant="italic">σ</mml:mi><mml:mi>n</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msubsup><mml:mo>(</mml:mo><mml:mi mathvariant="bold-italic">x</mml:mi><mml:mo>)</mml:mo><mml:mo>)</mml:mo><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

          where <inline-formula><mml:math id="M140" display="inline"><mml:mrow><mml:msubsup><mml:mi mathvariant="italic">σ</mml:mi><mml:mi>n</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msubsup><mml:mo>(</mml:mo><mml:mi mathvariant="bold-italic">x</mml:mi><mml:mo>)</mml:mo><mml:mo>=</mml:mo><mml:msup><mml:mi mathvariant="italic">α</mml:mi><mml:mrow><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow></mml:msup><mml:mo>+</mml:mo><mml:msup><mml:mi>x</mml:mi><mml:mi mathvariant="normal">⊤</mml:mi></mml:msup><mml:msub><mml:mi mathvariant="normal">Σ</mml:mi><mml:mi>n</mml:mi></mml:msub><mml:mi mathvariant="bold-italic">x</mml:mi></mml:mrow></mml:math></inline-formula> decomposes total predictive uncertainty into aleatoric noise <inline-formula><mml:math id="M141" display="inline"><mml:mrow><mml:msup><mml:mi mathvariant="italic">α</mml:mi><mml:mrow><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow></mml:msup></mml:mrow></mml:math></inline-formula> and parameter uncertainty encoded in the posterior covariance <inline-formula><mml:math id="M142" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="normal">Σ</mml:mi><mml:mi>n</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>. A separate BRR meta-learner is fitted independently for each forecasting horizon, receiving the two-dimensional out-of-fold prediction vector generated by the two base-learner instances trained for that horizon.</p>
</sec>
<sec id="Ch1.S2.SS5">
  <label>2.5</label><title>Multi-Horizon Forecasting Strategy</title>
      <p id="d2e3314">Multi-step forecasting can be approached through four principal strategies, including recursive, direct, hybrid, and multi-input multi-output (MIMO) approaches (Ben Taieb et al., 2012). In the present study, the direct method is adopted: a fully independent instance of the ensemble is trained separately for each forecasting horizon, with no parameter sharing across horizons. This choice avoids the error accumulation inherent to the recursive and hybrid strategies, and avoids the output-layer distribution mismatch that can arise when a single MIMO model must simultaneously minimise losses at horizons with substantially different error magnitudes. The additional training cost of maintaining three independent model instances is computationally manageable given the offline nature of model development for this application.</p>
</sec>
</sec>
<sec id="Ch1.S3">
  <label>3</label><title>Experiments</title>
<sec id="Ch1.S3.SS1">
  <label>3.1</label><title>Datasets and metrics</title>
<sec id="Ch1.S3.SS1.SSS1">
  <label>3.1.1</label><title>Site data</title>
      <p id="d2e3340">This study focuses on a road segment in the northwest inland plain of Jiangsu Province, China, adjacent to the Longhai Railway overpass. The region belongs to a warm temperate semi-humid monsoon climate zone, with minimum daily temperatures in winter reaching approximately <inline-formula><mml:math id="M143" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>3 °C. The terrain is predominantly flat, which facilitates the horizontal distribution of meteorological variables and minimizes orographic effects that would otherwise compromise the spatial representativeness of gridded reanalysis data.</p>
      <p id="d2e3350">The primary dataset was collected from road meteorological monitoring station M9393, located west of the Longhai Railway overpass at coordinates (34.30° N, 117.04° E), at a temporal resolution of five minutes. The dataset spans four consecutive winter periods: December 2020 to February 2021, December 2021 to February 2022, December 2022 to February 2023, and December 2023 to February 2024. Infrared remote sensing sensors recorded near-surface meteorological variables including visibility (<inline-formula><mml:math id="M144" display="inline"><mml:mi>V</mml:mi></mml:math></inline-formula>), air temperature (AT), relative humidity (RH), precipitation (<inline-formula><mml:math id="M145" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula>), wind speed (WS), wind direction (WD), and road surface temperature (RST), totaling 102 431 raw observations.</p>
      <p id="d2e3367">Data quality control was applied in several steps. First, the dataset was divided into training and testing subsets. The training subset comprised data from 1 December 2020 to 28 February 2023, and the testing subset comprised data from 1 December 2023 to 29 February  2024. Physically implausible values were then rejected as sensor errors: RST observations exceeding 50 °C or falling below <inline-formula><mml:math id="M146" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>40 °C, negative precipitation values, and relative humidity outside the range [0 %, 100 %] were removed. For gaps not exceeding one consecutive hour, linear interpolation is applied using the nearest valid observations bracketing the gap, drawn only from that same period. If the right-side interpolation anchor falls within the same natural hour as the missing slot(s), the resulting hourly mean incorporates no observation from any later hour, and no cross-hour information enters the model by construction.</p>
      <p id="d2e3377">Longer gaps, exceeding one consecutive hour and occurring primarily during scheduled sensor maintenance, were filled using the mean of the observed values recorded at the same hour-of-day and same calendar period across the three training winters, and applied identically regardless of whether the corresponding gap occurred in the training or test period; because this value is derived exclusively from training-period observations at matching calendar times from three prior winters, it introduces no look-ahead information from any test-period observation. This long-gap climatological fill accounts for 1.18 % of 5 min test-period samples at M9393. Following imputation, all variables were resampled to hourly resolution within each period: precipitation was aggregated as the hourly total, and all other variables were computed as hourly means to suppress transient sub-hourly fluctuations. The resulting hourly series exhibits a near-unity variance ratio relative to the native 5 min series and matching significant periodicities (Figs. S1–S2 and Table S1 in the Supplement). This yielded a final dataset of 8664 samples for model development.</p>
      <p id="d2e3381">RST exhibits significant diurnal periodicity and lag responses to meteorological forcing, reflecting the complex and nonlinear nature of heat exchange at the pavement–atmosphere interface (Cheng et al., 2021). Figure 5 illustrates the daily variation of RST during December 2020, showing a consistent diurnal cycle with a fluctuation range of approximately 5 to 15 °C and strong hour-to-hour autocorrelation across successive days. Figure 6 presents the mean 24 h diurnal RST profile and its 95 % confidence interval computed across the full study period. RST reaches its minimum during the pre-dawn hours, rises gradually through the morning, reaches a mean daily maximum of approximately 6 °C around 14:00, and subsequently declines through the evening. The width of the confidence interval reflects the day-to-day variability at each hour arising from meteorological disturbances superimposed on the underlying diurnal cycle. These characteristics motivate the inclusion of historical RST sequence an input feature.</p>

      <fig id="F5" specific-use="star"><label>Figure 5</label><caption><p id="d2e3386">Periodic variation of RST. <inline-formula><mml:math id="M147" display="inline"><mml:mi>T</mml:mi></mml:math></inline-formula> represents a period, with time corresponding to one day.</p></caption>
            <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f05.png"/>

          </fig>

      <fig id="F6" specific-use="star"><label>Figure 6</label><caption><p id="d2e3404">The 24 h diurnal variation curve of RST. The shaded blue part of the graph is the 95 % confidence interval for the RST.</p></caption>
            <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f06.png"/>

          </fig>

</sec>
<sec id="Ch1.S3.SS1.SSS2">
  <label>3.1.2</label><title>ERA5-Land reanalysis data</title>
      <p id="d2e3421">ERA5-Land reanalysis data were obtained for the same temporal coverage as the station observations to investigate the potential contribution of physically meaningful auxiliary variables. ERA5-Land is produced by the European Centre for Medium-Range Weather Forecasts (ECMWF) through land surface model simulations driven by ERA5 atmospheric forcing, and is publicly available from the Copernicus Climate Change Service (<uri>https://cds.climate.copernicus.eu/</uri>, last access: 18 April 2026). The dataset provides hourly estimates at a spatial resolution of 0.1<inline-formula><mml:math id="M148" display="inline"><mml:mrow><mml:mi mathvariant="italic">°</mml:mi><mml:mspace width="0.125em" linebreak="nobreak"/><mml:mo>×</mml:mo></mml:mrow></mml:math></inline-formula> 0.1° (approximately 9–11 km), representing area-averaged surface conditions rather than point-scale measurements.</p>
      <p id="d2e3438">Eight candidate variables with established thermodynamic relevance to road surface temperature dynamics were extracted as the initial variable pool: soil temperature (ST), surface net solar radiation (SSR), surface net thermal radiation (STR), surface sensible heat flux (SSHF), surface latent heat flux (SLHF), forecast albedo (FAL), evaporation from bare soil (EVABS), and total evaporation (<inline-formula><mml:math id="M149" display="inline"><mml:mi>E</mml:mi></mml:math></inline-formula>). These variables collectively represent the principal radiative, turbulent, and latent heat exchange terms in the surface energy balance, and their physical relevance to RST evolution is established in the literature (Chen et al., 2019; Hermansson, 2004). Variables ultimately retained for model input were determined through the correlation screening procedure described in Sect. 3.2.</p>
      <p id="d2e3448">Spatial matching was performed using the nearest-neighbour method, identifying the grid point (34.30° N, 117.00° E) closest to the station. This approach is considered appropriate given the flat terrain of the study region, which minimizes the representativeness error associated with nearest-neighbour interpolation (Muñoz-Sabater et al., 2021). All variables were temporally aligned with the station dataset by converting from UTC to China Standard Time (UTC<inline-formula><mml:math id="M150" display="inline"><mml:mo>+</mml:mo></mml:math></inline-formula>8). Unit conversions were applied where necessary: SSR, STR, SSHF, and SLHF, provided as hourly accumulated values in J m<sup>−2</sup>, were divided by 3600 to obtain mean flux densities in W m<sup>−2</sup>; ST was converted from Kelvin to degrees Celsius.</p>

<table-wrap id="T2" specific-use="star"><label>Table 2</label><caption><p id="d2e3486">Summary of the dataset.</p></caption><oasis:table frame="topbot"><oasis:tgroup cols="6">
     <oasis:colspec colnum="1" colname="col1" align="left"/>
     <oasis:colspec colnum="2" colname="col2" align="right"/>
     <oasis:colspec colnum="3" colname="col3" align="right"/>
     <oasis:colspec colnum="4" colname="col4" align="right"/>
     <oasis:colspec colnum="5" colname="col5" align="right"/>
     <oasis:colspec colnum="6" colname="col6" align="left"/>
     <oasis:thead>
       <oasis:row rowsep="1">
         <oasis:entry colname="col1">Feature</oasis:entry>
         <oasis:entry colname="col2">Mean</oasis:entry>
         <oasis:entry colname="col3">Std</oasis:entry>
         <oasis:entry colname="col4">Min</oasis:entry>
         <oasis:entry colname="col5">Max</oasis:entry>
         <oasis:entry colname="col6">Unit</oasis:entry>
       </oasis:row>
     </oasis:thead>
     <oasis:tbody>
       <oasis:row>
         <oasis:entry colname="col1">Visibility</oasis:entry>
         <oasis:entry colname="col2">7942.28</oasis:entry>
         <oasis:entry colname="col3">3355.89</oasis:entry>
         <oasis:entry colname="col4">56.00</oasis:entry>
         <oasis:entry colname="col5">30 000.00</oasis:entry>
         <oasis:entry colname="col6">m</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Air temperature</oasis:entry>
         <oasis:entry colname="col2">2.08</oasis:entry>
         <oasis:entry colname="col3">5.29</oasis:entry>
         <oasis:entry colname="col4"><inline-formula><mml:math id="M153" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>15.15</oasis:entry>
         <oasis:entry colname="col5">23.79</oasis:entry>
         <oasis:entry colname="col6">°C</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Relative humidity</oasis:entry>
         <oasis:entry colname="col2">55.97</oasis:entry>
         <oasis:entry colname="col3">24.14</oasis:entry>
         <oasis:entry colname="col4">1.76</oasis:entry>
         <oasis:entry colname="col5">100.00</oasis:entry>
         <oasis:entry colname="col6">%</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Precipitation</oasis:entry>
         <oasis:entry colname="col2">0.01</oasis:entry>
         <oasis:entry colname="col3">0.12</oasis:entry>
         <oasis:entry colname="col4">0.00</oasis:entry>
         <oasis:entry colname="col5">4.60</oasis:entry>
         <oasis:entry colname="col6">mm</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Wind speed</oasis:entry>
         <oasis:entry colname="col2">1.89</oasis:entry>
         <oasis:entry colname="col3">1.19</oasis:entry>
         <oasis:entry colname="col4">0.30</oasis:entry>
         <oasis:entry colname="col5">8.72</oasis:entry>
         <oasis:entry colname="col6">m s<sup>−1</sup></oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Wind direction</oasis:entry>
         <oasis:entry colname="col2">175.41</oasis:entry>
         <oasis:entry colname="col3">92.87</oasis:entry>
         <oasis:entry colname="col4">0.03</oasis:entry>
         <oasis:entry colname="col5">359.95</oasis:entry>
         <oasis:entry colname="col6">°</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Road surface temperature</oasis:entry>
         <oasis:entry colname="col2">3.72</oasis:entry>
         <oasis:entry colname="col3">5.59</oasis:entry>
         <oasis:entry colname="col4"><inline-formula><mml:math id="M155" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>12.99</oasis:entry>
         <oasis:entry colname="col5">26.52</oasis:entry>
         <oasis:entry colname="col6">°C</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Soil temperature</oasis:entry>
         <oasis:entry colname="col2">3.83</oasis:entry>
         <oasis:entry colname="col3">3.92</oasis:entry>
         <oasis:entry colname="col4"><inline-formula><mml:math id="M156" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>2.60</oasis:entry>
         <oasis:entry colname="col5">20.88</oasis:entry>
         <oasis:entry colname="col6">°C</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Surface net solar radiation</oasis:entry>
         <oasis:entry colname="col2">99.42</oasis:entry>
         <oasis:entry colname="col3">164.05</oasis:entry>
         <oasis:entry colname="col4">0.00</oasis:entry>
         <oasis:entry colname="col5">652.08</oasis:entry>
         <oasis:entry colname="col6">W m<sup>−2</sup></oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Surface net thermal radiation</oasis:entry>
         <oasis:entry colname="col2">68.53</oasis:entry>
         <oasis:entry colname="col3">348.85</oasis:entry>
         <oasis:entry colname="col4">0.00</oasis:entry>
         <oasis:entry colname="col5">2601.48</oasis:entry>
         <oasis:entry colname="col6">W m<sup>−2</sup></oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Surface sensible heat flux</oasis:entry>
         <oasis:entry colname="col2">30.95</oasis:entry>
         <oasis:entry colname="col3">105.62</oasis:entry>
         <oasis:entry colname="col4">0.00</oasis:entry>
         <oasis:entry colname="col5">1671.75</oasis:entry>
         <oasis:entry colname="col6">W m<sup>−2</sup></oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Surface latent heat flux</oasis:entry>
         <oasis:entry colname="col2">20.65</oasis:entry>
         <oasis:entry colname="col3">109.59</oasis:entry>
         <oasis:entry colname="col4">0.00</oasis:entry>
         <oasis:entry colname="col5">1522.77</oasis:entry>
         <oasis:entry colname="col6">W m<sup>−2</sup></oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Forecast albedo</oasis:entry>
         <oasis:entry colname="col2">0.17</oasis:entry>
         <oasis:entry colname="col3">0.05</oasis:entry>
         <oasis:entry colname="col4">0.15</oasis:entry>
         <oasis:entry colname="col5">0.59</oasis:entry>
         <oasis:entry colname="col6">–</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Evaporation from bare soil</oasis:entry>
         <oasis:entry colname="col2"><inline-formula><mml:math id="M161" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>0.42</oasis:entry>
         <oasis:entry colname="col3">0.32</oasis:entry>
         <oasis:entry colname="col4"><inline-formula><mml:math id="M162" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>1.70</oasis:entry>
         <oasis:entry colname="col5">0.00</oasis:entry>
         <oasis:entry colname="col6">mm</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Total evaporation</oasis:entry>
         <oasis:entry colname="col2"><inline-formula><mml:math id="M163" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>0.56</oasis:entry>
         <oasis:entry colname="col3">0.37</oasis:entry>
         <oasis:entry colname="col4"><inline-formula><mml:math id="M164" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>2.25</oasis:entry>
         <oasis:entry colname="col5">0.00</oasis:entry>
         <oasis:entry colname="col6">mm</oasis:entry>
       </oasis:row>
     </oasis:tbody>
   </oasis:tgroup></oasis:table></table-wrap>

</sec>
<sec id="Ch1.S3.SS1.SSS3">
  <label>3.1.3</label><title>Evaluation index</title>
      <p id="d2e3972">In this study, multiple evaluation metrics are employed to comprehensively evaluate the predictive performance of the models, including the mean absolute error (MAE), root mean squared error (RMSE), symmetric mean absolute percentage error (sMAPE), and coefficient of determination (<inline-formula><mml:math id="M165" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup></mml:mrow></mml:math></inline-formula>).

                  <disp-formula specific-use="gather" content-type="numbered"><mml:math id="M166" display="block"><mml:mtable displaystyle="true"><mml:mlabeledtr id="Ch1.E21"><mml:mtd><mml:mtext>21</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle class="stylechange" displaystyle="true"/><mml:mi mathvariant="normal">MAE</mml:mi><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mn mathvariant="normal">1</mml:mn><mml:mi>n</mml:mi></mml:mfrac></mml:mstyle><mml:msubsup><mml:mo>∑</mml:mo><mml:mrow><mml:mi>i</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>n</mml:mi></mml:msubsup><mml:mfenced close="|" open="|"><mml:mrow><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo>-</mml:mo><mml:msub><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo mathvariant="normal" stretchy="false">^</mml:mo></mml:mover><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:mfenced><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr><mml:mlabeledtr id="Ch1.E22"><mml:mtd><mml:mtext>22</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle displaystyle="true" class="stylechange"/><mml:mi mathvariant="normal">RMSE</mml:mi><mml:mo>=</mml:mo><mml:msqrt><mml:mrow><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mn mathvariant="normal">1</mml:mn><mml:mi>n</mml:mi></mml:mfrac></mml:mstyle><mml:msubsup><mml:mo>∑</mml:mo><mml:mrow><mml:mi>i</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>n</mml:mi></mml:msubsup><mml:mo>(</mml:mo><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo>-</mml:mo><mml:msub><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo mathvariant="normal" stretchy="false">^</mml:mo></mml:mover><mml:mi>i</mml:mi></mml:msub><mml:msup><mml:mo>)</mml:mo><mml:mn mathvariant="normal">2</mml:mn></mml:msup></mml:mrow></mml:msqrt><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr><mml:mlabeledtr id="Ch1.E23"><mml:mtd><mml:mtext>23</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle displaystyle="true" class="stylechange"/><mml:mi mathvariant="normal">sMAPE</mml:mi><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mn mathvariant="normal">1</mml:mn><mml:mi>n</mml:mi></mml:mfrac></mml:mstyle><mml:msubsup><mml:mo>∑</mml:mo><mml:mrow><mml:mi>i</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>n</mml:mi></mml:msubsup><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mrow><mml:mn mathvariant="normal">2</mml:mn><mml:mfenced open="|" close="|"><mml:mrow><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo>-</mml:mo><mml:msub><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo stretchy="false" mathvariant="normal">^</mml:mo></mml:mover><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:mfenced></mml:mrow><mml:mrow><mml:mfenced close="|" open="|"><mml:mrow><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:mfenced><mml:mo>+</mml:mo><mml:mfenced close="|" open="|"><mml:mrow><mml:msub><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo mathvariant="normal" stretchy="false">^</mml:mo></mml:mover><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:mfenced><mml:mo>+</mml:mo><mml:mi mathvariant="italic">ϵ</mml:mi></mml:mrow></mml:mfrac></mml:mstyle><mml:mo>×</mml:mo><mml:mn mathvariant="normal">100</mml:mn><mml:mspace width="0.125em" linebreak="nobreak"/><mml:mi mathvariant="italic">%</mml:mi><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr><mml:mlabeledtr id="Ch1.E24"><mml:mtd><mml:mtext>24</mml:mtext></mml:mtd><mml:mtd><mml:mrow><mml:mstyle class="stylechange" displaystyle="true"/><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn><mml:mo>-</mml:mo><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mrow><mml:msubsup><mml:mo>∑</mml:mo><mml:mrow><mml:mi>i</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>n</mml:mi></mml:msubsup><mml:mo>(</mml:mo><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo>-</mml:mo><mml:msub><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo mathvariant="normal" stretchy="false">^</mml:mo></mml:mover><mml:mi>i</mml:mi></mml:msub><mml:msup><mml:mo>)</mml:mo><mml:mn mathvariant="normal">2</mml:mn></mml:msup></mml:mrow><mml:mrow><mml:msubsup><mml:mo>∑</mml:mo><mml:mrow><mml:mi>i</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow><mml:mi>n</mml:mi></mml:msubsup><mml:mo>(</mml:mo><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub><mml:mo>-</mml:mo><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo mathvariant="normal">‾</mml:mo></mml:mover><mml:msup><mml:mo>)</mml:mo><mml:mn mathvariant="normal">2</mml:mn></mml:msup></mml:mrow></mml:mfrac></mml:mstyle><mml:mo>,</mml:mo></mml:mrow></mml:mtd></mml:mlabeledtr></mml:mtable></mml:math></disp-formula>

            where <inline-formula><mml:math id="M167" display="inline"><mml:mrow><mml:msub><mml:mi>y</mml:mi><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the observed value, <inline-formula><mml:math id="M168" display="inline"><mml:mrow><mml:msub><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo stretchy="false" mathvariant="normal">^</mml:mo></mml:mover><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the predicted value, <inline-formula><mml:math id="M169" display="inline"><mml:mover accent="true"><mml:mi>y</mml:mi><mml:mo mathvariant="normal">‾</mml:mo></mml:mover></mml:math></inline-formula> is the mean of the observed values, and <inline-formula><mml:math id="M170" display="inline"><mml:mi>n</mml:mi></mml:math></inline-formula> is the total number of samples. and <inline-formula><mml:math id="M171" display="inline"><mml:mi mathvariant="italic">ϵ</mml:mi></mml:math></inline-formula> is a small constant (<inline-formula><mml:math id="M172" display="inline"><mml:mrow><mml:mi mathvariant="italic">ϵ</mml:mi><mml:mo>=</mml:mo><mml:msup><mml:mn mathvariant="normal">10</mml:mn><mml:mrow><mml:mo>-</mml:mo><mml:mn mathvariant="normal">8</mml:mn></mml:mrow></mml:msup></mml:mrow></mml:math></inline-formula>) introduced to prevent division by zero when both the observed and predicted values simultaneously approach zero. All reported evaluation metrics were computed exclusively against directly observed RST values; imputed hourly values were retained as model inputs where they occur but were excluded from the evaluation targets.</p>
</sec>
</sec>
<sec id="Ch1.S3.SS2">
  <label>3.2</label><title>Feature extraction</title>
      <p id="d2e4345">To select a parsimonious and informative input feature set, Spearman rank correlation coefficients (Spearman, 1961) were computed between each candidate variable and winter RST. This nonparametric measure is used in preference to Pearson correlation given the potential for nonlinear monotonic relationships between meteorological variables and RST. The coefficient is calculated as:

            <disp-formula id="Ch1.E25" content-type="numbered"><label>25</label><mml:math id="M173" display="block"><mml:mrow><mml:mi mathvariant="italic">ρ</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">1</mml:mn><mml:mo>-</mml:mo><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mrow><mml:mn mathvariant="normal">6</mml:mn><mml:mi mathvariant="normal">Σ</mml:mi><mml:mspace width="0.125em" linebreak="nobreak"/><mml:msubsup><mml:mi mathvariant="normal">d</mml:mi><mml:mi>i</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msubsup></mml:mrow><mml:mrow><mml:mi>n</mml:mi><mml:mfenced close=")" open="("><mml:mrow><mml:msup><mml:mi>n</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup><mml:mo>-</mml:mo><mml:mn mathvariant="normal">1</mml:mn></mml:mrow></mml:mfenced></mml:mrow></mml:mfrac></mml:mstyle><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

          where <inline-formula><mml:math id="M174" display="inline"><mml:mrow><mml:msub><mml:mi>d</mml:mi><mml:mi>i</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> is the rank difference between paired observations, and <inline-formula><mml:math id="M175" display="inline"><mml:mi>n</mml:mi></mml:math></inline-formula> is the number of samples.</p>
      <p id="d2e4411">As illustrated in Fig. 7, among the station-observed variables, air temperature exhibits the strongest positive correlation with RST (<inline-formula><mml:math id="M176" display="inline"><mml:mrow><mml:mi mathvariant="italic">ρ</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">0.93</mml:mn></mml:mrow></mml:math></inline-formula>), consistent with its recognised role as the primary meteorological driver of near-surface pavement temperature (Chen et al., 2019). Wind speed and precipitation display moderate positive correlations of 0.29 and 0.27, respectively, while relative humidity shows a moderate negative correlation (<inline-formula><mml:math id="M177" display="inline"><mml:mrow><mml:mi mathvariant="italic">ρ</mml:mi><mml:mo>=</mml:mo></mml:mrow></mml:math></inline-formula> <inline-formula><mml:math id="M178" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>0.25), consistent with the known cooling effect of high-humidity conditions at the pavement surface (Gui et al., 2007). Visibility and wind direction exhibit correlations below 0.20 and are excluded from the input feature set. It is additionally noted that wind direction is a circular variable for which rank-based correlation computed on raw degree values can be misleading; however, the low correlation magnitude observed here is consistent across both conventional and circular-adapted metrics, and the exclusion of wind direction from the final feature set is supported regardless of the association measure applied. The RST sequence is retained as an input variable because, as a directly measured surface quantity, it reflects the integrated effect of prior meteorological forcing at this monitored cross-section and provides information that is complementary to the instantaneous meteorological observations. Based on these findings, the input feature set for the primary analytical perspectives comprises five station-observed variables: AT, RH, <inline-formula><mml:math id="M179" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula>, WS, and historical RST.</p>

      <fig id="F7" specific-use="star"><label>Figure 7</label><caption><p id="d2e4452">Spearman correlation coefficient plot between RST and various meteorological factors and ERA5_Land physical factors.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f07.png"/>

        </fig>

      <p id="d2e4462">For the analytical perspective that investigates the added predictive value of ERA5-Land reanalysis variables (Sect. 4.2, Variable combination 2), the correlation analysis is extended to the full candidate variable pool including both station observations and ERA5-Land fields. Among the ERA5-Land variables, soil temperature shows the strongest positive correlation with RST (<inline-formula><mml:math id="M180" display="inline"><mml:mrow><mml:mi mathvariant="italic">ρ</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">0.49</mml:mn></mml:mrow></mml:math></inline-formula>), which may reflect its role as an integrative indicator of subsurface and near-surface thermal conditions. Surface net solar radiation and forecast albedo exhibit correlations of 0.24 and <inline-formula><mml:math id="M181" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>0.20, respectively, consistent with the known sensitivity of pavement temperature to radiative forcing (Qin et al., 2022). Evaporation from bare soil shows a moderate negative correlation (<inline-formula><mml:math id="M182" display="inline"><mml:mrow><mml:mi mathvariant="italic">ρ</mml:mi><mml:mo>=</mml:mo></mml:mrow></mml:math></inline-formula> <inline-formula><mml:math id="M183" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>0.23), consistent with the cooling effect associated with surface moisture conditions. The remaining ERA5-Land variables exhibit correlations below 0.20 with RST and are excluded. Accordingly, the augmented input set for Variable combination 2 in Sect. 4.2 comprises nine variables: AT, RH, <inline-formula><mml:math id="M184" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula>, WS, ST, SSR, FAL, EVABS, and historical RST.</p>
</sec>
<sec id="Ch1.S3.SS3">
  <label>3.3</label><title>Optimisation</title>
<sec id="Ch1.S3.SS3.SSS1">
  <label>3.3.1</label><title>Sliding window size</title>
      <p id="d2e4523">The sliding window method is employed to construct input variables and prediction targets for the model, and the selection of the window size directly influences the ability of the model to capture temperature periodicity and lag effects of meteorological variables. To identify the optimal window size, we conducted a comparative analysis of the improved base models under varying window configurations (Table 3). This analysis revealed that a 24 h sliding window offers an effective balance between predictive performance and computational efficiency, while also matching the timescale of the diurnal RST cycle, enabling the model to learn the recurring day–night thermal rhythm directly from the historical RST and meteorological sequence at this site. Consequently, the 24 h window was selected as the standardized configuration for all subsequent modelling experiments.</p>

<table-wrap id="T3" specific-use="star"><label>Table 3</label><caption><p id="d2e4529">Impact of sliding window size on errors.</p></caption><oasis:table frame="topbot"><oasis:tgroup cols="7">
     <oasis:colspec colnum="1" colname="col1" align="left"/>
     <oasis:colspec colnum="2" colname="col2" align="center"/>
     <oasis:colspec colnum="3" colname="col3" align="center"/>
     <oasis:colspec colnum="4" colname="col4" align="center" colsep="1"/>
     <oasis:colspec colnum="5" colname="col5" align="center"/>
     <oasis:colspec colnum="6" colname="col6" align="center"/>
     <oasis:colspec colnum="7" colname="col7" align="center"/>
     <oasis:thead>
       <oasis:row>

         <oasis:entry rowsep="1" colname="col1" morerows="1">Window size</oasis:entry>

         <oasis:entry rowsep="1" namest="col2" nameend="col4" colsep="1">KNN-LSTM </oasis:entry>

         <oasis:entry rowsep="1" namest="col5" nameend="col7">BiLSTM-MHA </oasis:entry>

       </oasis:row>
       <oasis:row rowsep="1">

         <oasis:entry colname="col2">MAE (°C)</oasis:entry>

         <oasis:entry colname="col3">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col4">sMAPE (%)</oasis:entry>

         <oasis:entry colname="col5">MAE (°C)</oasis:entry>

         <oasis:entry colname="col6">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col7">sMAPE (%)</oasis:entry>

       </oasis:row>
     </oasis:thead>
     <oasis:tbody>
       <oasis:row>

         <oasis:entry colname="col1">6 h</oasis:entry>

         <oasis:entry colname="col2">0.507</oasis:entry>

         <oasis:entry colname="col3">0.665</oasis:entry>

         <oasis:entry colname="col4">27.577</oasis:entry>

         <oasis:entry colname="col5">0.487</oasis:entry>

         <oasis:entry colname="col6">0.655</oasis:entry>

         <oasis:entry colname="col7">25.871</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">12 h</oasis:entry>

         <oasis:entry colname="col2">0.517</oasis:entry>

         <oasis:entry colname="col3">0.669</oasis:entry>

         <oasis:entry colname="col4">28.550</oasis:entry>

         <oasis:entry colname="col5">0.431</oasis:entry>

         <oasis:entry colname="col6">0.617</oasis:entry>

         <oasis:entry colname="col7">22.000</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">24 h</oasis:entry>

         <oasis:entry colname="col2">0.389</oasis:entry>

         <oasis:entry colname="col3">0.536</oasis:entry>

         <oasis:entry colname="col4">21.348</oasis:entry>

         <oasis:entry colname="col5">0.393</oasis:entry>

         <oasis:entry colname="col6">0.545</oasis:entry>

         <oasis:entry colname="col7">20.783</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">48 h</oasis:entry>

         <oasis:entry colname="col2">0.414</oasis:entry>

         <oasis:entry colname="col3">0.575</oasis:entry>

         <oasis:entry colname="col4">22.357</oasis:entry>

         <oasis:entry colname="col5">0.662</oasis:entry>

         <oasis:entry colname="col6">0.938</oasis:entry>

         <oasis:entry colname="col7">30.120</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">72 h</oasis:entry>

         <oasis:entry colname="col2">0.399</oasis:entry>

         <oasis:entry colname="col3">0.540</oasis:entry>

         <oasis:entry colname="col4">22.139</oasis:entry>

         <oasis:entry colname="col5">0.415</oasis:entry>

         <oasis:entry colname="col6">0.569</oasis:entry>

         <oasis:entry colname="col7">22.772</oasis:entry>

       </oasis:row>
     </oasis:tbody>
   </oasis:tgroup></oasis:table></table-wrap>

</sec>
<sec id="Ch1.S3.SS3.SSS2">
  <label>3.3.2</label><title>Training</title>
      <p id="d2e4725">Prior to model training, all variables were standardized using <inline-formula><mml:math id="M185" display="inline"><mml:mi>Z</mml:mi></mml:math></inline-formula>-score normalisation to accelerate convergence and eliminate the influence of differing measurement scales. Normalisation parameters were estimated exclusively from the training set to prevent information leakage:

              <disp-formula id="Ch1.E26" content-type="numbered"><label>26</label><mml:math id="M186" display="block"><mml:mrow><mml:msub><mml:mi>X</mml:mi><mml:mi mathvariant="normal">scaled</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:mstyle displaystyle="true"><mml:mfrac style="display"><mml:mrow><mml:msub><mml:mi>X</mml:mi><mml:mi mathvariant="normal">train</mml:mi></mml:msub><mml:mo>-</mml:mo><mml:msub><mml:mi mathvariant="italic">μ</mml:mi><mml:mi mathvariant="normal">train</mml:mi></mml:msub></mml:mrow><mml:mrow><mml:msub><mml:mi mathvariant="italic">σ</mml:mi><mml:mi mathvariant="normal">train</mml:mi></mml:msub></mml:mrow></mml:mfrac></mml:mstyle><mml:mo>,</mml:mo></mml:mrow></mml:math></disp-formula>

            where <inline-formula><mml:math id="M187" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="italic">μ</mml:mi><mml:mi mathvariant="normal">train</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> and <inline-formula><mml:math id="M188" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="italic">σ</mml:mi><mml:mi mathvariant="normal">train</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> are the mean and standard deviation computed from the training set only. Model outputs were inverse-transformed to degrees Celsius prior to evaluation. Accordingly, all reported MAE and RMSE values are in °C.</p>
      <p id="d2e4793">To assess the climatological representativeness of the test period, the mean air temperature and RST during the test winter were compared against the three-winter training period mean. The test-period mean air temperature (<inline-formula><mml:math id="M189" display="inline"><mml:mo lspace="0mm">-</mml:mo></mml:math></inline-formula>1.29 °C) and mean RST (0.49 °C) deviate by less than 1.0 °C from the corresponding training-period means (<inline-formula><mml:math id="M190" display="inline"><mml:mo lspace="0mm">-</mml:mo></mml:math></inline-formula>0.31 °C and 1.40 °C, respectively), and the proportions of sub-zero RST hours are comparable (45.74 % and 47.53 %). These statistics confirm that the test winter is climatologically representative of the study period and that the evaluation is not materially affected by anomalous thermal conditions.</p>
      <p id="d2e4810">In the KNN-LSTM model, the hyperparameter <inline-formula><mml:math id="M191" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> specifies the number of nearest-neighbour historical sequences used to construct the similarity feature matrix <inline-formula><mml:math id="M192" display="inline"><mml:mrow><mml:msub><mml:mi mathvariant="bold">F</mml:mi><mml:mi>t</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>, and directly governs the trade-off between local pattern specificity and representational stability. A value of <inline-formula><mml:math id="M193" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> that is too small risks overfitting to unrepresentative neighbours, while an excessively large <inline-formula><mml:math id="M194" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> may dilute the local similarity signal by averaging over dissimilar historical states. To determine the optimal <inline-formula><mml:math id="M195" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula>, a grid search over <inline-formula><mml:math id="M196" display="inline"><mml:mrow><mml:mi>K</mml:mi><mml:mo>∈</mml:mo><mml:mo mathvariant="italic">{</mml:mo><mml:mn mathvariant="normal">3</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">5</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">7</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">9</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">11</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">13</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">15</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">17</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">19</mml:mn><mml:mo mathvariant="italic">}</mml:mo></mml:mrow></mml:math></inline-formula> was conducted by evaluating MAE and RMSE on the validation set. As shown in Fig. 8, both metrics decrease monotonically with increasing <inline-formula><mml:math id="M197" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula>, reaching their minimum at <inline-formula><mml:math id="M198" display="inline"><mml:mrow><mml:mi>K</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">15</mml:mn></mml:mrow></mml:math></inline-formula>, beyond which performance stabilizes. Accordingly, <inline-formula><mml:math id="M199" display="inline"><mml:mrow><mml:mi>K</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">15</mml:mn></mml:mrow></mml:math></inline-formula> is adopted as the fixed hyperparameter for all subsequent experiments. It should be noted that this search is a one-time offline procedure performed during model development; during inference, the KNN-LSTM model executes a single forward pass using the fixed <inline-formula><mml:math id="M200" display="inline"><mml:mrow><mml:mi>K</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">15</mml:mn></mml:mrow></mml:math></inline-formula>, incurring no additional computational overhead.</p>

      <fig id="F8" specific-use="star"><label>Figure 8</label><caption><p id="d2e4948">The effect of <inline-formula><mml:math id="M201" display="inline"><mml:mi>K</mml:mi></mml:math></inline-formula> value on MAE and RMSE.</p></caption>
            <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f08.png"/>

          </fig>

      <p id="d2e4964">The KNN-LSTM model employs a three-layer LSTM architecture with 64, 32, and 16 hidden units in successive layers, with tanh activation in the first layer and ReLU in deeper layers to mitigate the vanishing gradient problem (Glorot et al., 2011). The BiLSTM-MHA model uses a bidirectional LSTM with 128 units per direction, followed by a multi-head attention mechanism with 4 heads and key dimension 64, layer normalisation with residual connection, global average pooling, and a single-unit dense output layer. Both models are trained with the Adam optimizer (Kingma and Ba, 2014) at a learning rate of 0.001, mean squared error loss, batch size 32, a maximum of 100 epochs, and early stopping with patience 20 to prevent overfitting. L2 regularization is applied to all LSTM layers to further improve generalization. This architecture and training configuration is shared across the three forecasting horizons; for each horizon, a separate model instance with this identical configuration is trained independently, yielding horizon-specific weights with no parameter sharing across horizons.</p>
</sec>
</sec>
<sec id="Ch1.S3.SS4">
  <label>3.4</label><title>Complementarity and ablation analysis</title>
      <p id="d2e4976">The effectiveness of a stacking ensemble depends critically on base learner diversity: models that commit errors in different input regions enable the meta-learner to achieve superior performance beyond any individual component. Computational constraints limit the number of base learners, as each requires independent out-of-fold training across all folds. KNN-LSTM and BiLSTM-MHA are selected for their partially complementary predictive characteristics with respect to winter RST spatiotemporal structure. KNN-LSTM integrates similarity-based feature augmentation with sequential modelling to identify local pattern recurrence through retrieval of analogous historical meteorological states, making it effective at capturing regime-specific behaviours tied to locally recurring weather conditions. BiLSTM-MHA employs bidirectional recurrent processing with multi-head self-attention to dynamically weight temporal features across the full input window, demonstrating superior capability in extracting global temporal dependencies and long-range patterns during thermally complex periods involving sustained trends or abrupt transitions.</p>
      <p id="d2e4980">To empirically validate this complementarity prior to ensemble construction, we conduct a complementarity analysis and an ablation study, both at the 1 h forecasting horizon where the ensemble signal is cleanest. Results are presented in Fig. 9 and Table 4. The Pearson correlation coefficient between the residual series of the two base learners is <inline-formula><mml:math id="M202" display="inline"><mml:mrow><mml:mi>r</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">0.889</mml:mn></mml:mrow></mml:math></inline-formula>, indicating substantial but imperfect co-variation in prediction errors. However, a high global correlation does not preclude conditional complementarity across distinct operating regimes (Kuncheva and Whitaker, 2003). As shown in Fig. 9b, the two models exhibit asymmetric error patterns across temperature regimes: under sub-zero conditions, BiLSTM-MHA achieves a lower MAE, whereas under above-zero conditions KNN-LSTM is more accurate. This regime-dependent asymmetry is noteworthy because sub-zero and above-zero conditions are typically associated with distinct meteorological forcing characteristics – sustained radiative cooling under clear skies versus more variable radiation-driven heating, respectively (Hermansson, 2004; Qin et al., 2022) – suggesting that the two architectures may be exploiting different statistical regularities in the RST signal. Figure 9c shows comparable daytime and nighttime performance at the aggregate level, indicating that temperature-regime complementarity, rather than the diurnal cycle per se, is the dominant source of predictive diversity between the two base learners. Figure 9d further confirms that neither model achieves consistent sample-level dominance, demonstrating sustained bidirectional predictive diversity that the meta-learner can systematically exploit.</p>

      <fig id="F9" specific-use="star"><label>Figure 9</label><caption><p id="d2e4997">Complementarity analysis of base learners: <bold>(a)</bold> residual correlation analysis, <bold>(b)</bold> temperature interval error decomposition, <bold>(c)</bold> day and night time error decomposition, and <bold>(d)</bold> time series advantage distribution.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f09.png"/>

        </fig>

      <p id="d2e5019">Three ablation configurations are evaluated to quantify the independent contribution of each architectural innovation (Table 4). Config-1 (LSTM <inline-formula><mml:math id="M203" display="inline"><mml:mo>+</mml:mo></mml:math></inline-formula> BiLSTM) establishes the baseline ensemble without any proposed innovations; Config-2 (LSTM <inline-formula><mml:math id="M204" display="inline"><mml:mo>+</mml:mo></mml:math></inline-formula> BiLSTM-MHA) isolates the contribution of multi-head attention by introducing the MHA mechanism to the second base learner while retaining the standard LSTM as the first; Config-3 (KNN-LSTM <inline-formula><mml:math id="M205" display="inline"><mml:mo>+</mml:mo></mml:math></inline-formula> BiLSTM) isolates the contribution of KNN-based similarity augmentation by replacing the first base learner with KNN-LSTM while retaining the standard BiLSTM as the second; and the full ILES (KNN-LSTM <inline-formula><mml:math id="M206" display="inline"><mml:mo>+</mml:mo></mml:math></inline-formula> BiLSTM-MHA) achieves the best performance across all metrics. Relative to Config-1, introducing multi-head attention alone reduces MAE by 0.035 °C (8.20 %), whereas introducing KNN-based similarity augmentation alone reduces MAE by 0.044 °C (10.30 %). Yet the combined gain of ILES (0.054 °C) is smaller than the arithmetic sum of the two individual gains (0.079 °C), indicating that KNN-LSTM and BiLSTM-MHA partially share their areas of improvement; nonetheless, the full ILES framework achieves the best overall performance, confirming that the two components provide partially complementary rather than redundant or synergistic contributions.</p>

<table-wrap id="T4" specific-use="star"><label>Table 4</label><caption><p id="d2e5053">Ablation study on the effect of multi-head attention and KNN-based similarity augmentation.</p></caption><oasis:table frame="topbot"><oasis:tgroup cols="6">
     <oasis:colspec colnum="1" colname="col1" align="left"/>
     <oasis:colspec colnum="2" colname="col2" align="left"/>
     <oasis:colspec colnum="3" colname="col3" align="left"/>
     <oasis:colspec colnum="4" colname="col4" align="right"/>
     <oasis:colspec colnum="5" colname="col5" align="right"/>
     <oasis:colspec colnum="6" colname="col6" align="right"/>
     <oasis:thead>
       <oasis:row rowsep="1">
         <oasis:entry colname="col1">Configuration</oasis:entry>
         <oasis:entry colname="col2">Base learner 1</oasis:entry>
         <oasis:entry colname="col3">Base learner 2</oasis:entry>
         <oasis:entry colname="col4">MAE (°C)</oasis:entry>
         <oasis:entry colname="col5">RMSE (°C)</oasis:entry>
         <oasis:entry colname="col6">sMAPE (%)</oasis:entry>
       </oasis:row>
     </oasis:thead>
     <oasis:tbody>
       <oasis:row>
         <oasis:entry colname="col1">Config-1</oasis:entry>
         <oasis:entry colname="col2">LSTM</oasis:entry>
         <oasis:entry colname="col3">BiLSTM</oasis:entry>
         <oasis:entry colname="col4">0.427</oasis:entry>
         <oasis:entry colname="col5">0.597</oasis:entry>
         <oasis:entry colname="col6">21.842</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Config-2</oasis:entry>
         <oasis:entry colname="col2">LSTM</oasis:entry>
         <oasis:entry colname="col3">BiLSTM-MHA</oasis:entry>
         <oasis:entry colname="col4">0.392</oasis:entry>
         <oasis:entry colname="col5">0.552</oasis:entry>
         <oasis:entry colname="col6">21.033</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">Config-3</oasis:entry>
         <oasis:entry colname="col2">KNN-LSTM</oasis:entry>
         <oasis:entry colname="col3">BiLSTM</oasis:entry>
         <oasis:entry colname="col4">0.383</oasis:entry>
         <oasis:entry colname="col5">0.526</oasis:entry>
         <oasis:entry colname="col6">20.546</oasis:entry>
       </oasis:row>
       <oasis:row>
         <oasis:entry colname="col1">ILES</oasis:entry>
         <oasis:entry colname="col2">KNN-LSTM</oasis:entry>
         <oasis:entry colname="col3">BiLSTM-MHA</oasis:entry>
         <oasis:entry colname="col4">0.373</oasis:entry>
         <oasis:entry colname="col5">0.521</oasis:entry>
         <oasis:entry colname="col6">20.328</oasis:entry>
       </oasis:row>
     </oasis:tbody>
   </oasis:tgroup></oasis:table></table-wrap>

</sec>
<sec id="Ch1.S3.SS5">
  <label>3.5</label><title>Prediction performance evaluation</title>
      <p id="d2e5199">To systematically evaluate the proposed ILES framework and address the four objectives stated in Sect. 1, four analytical perspectives are designed as follows. <list list-type="order"><list-item>
      <p id="d2e5204"><italic>Comprehensive performance benchmarking.</italic> ILES is evaluated against ten models spanning three categories: a naive persistence baseline, nonlinear regression, traditional machine learning methods (Random Forest, XGBoost), standard deep learning architectures (GRU, LSTM, BiLSTM, CNN-LSTM), and the two proposed base learners (KNN-LSTM, BiLSTM-MHA). All models receive identical station-only inputs and are evaluated at 1, 3, and 6 h forecasting horizons using MAE, RMSE, sMAPE, and <inline-formula><mml:math id="M207" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup></mml:mrow></mml:math></inline-formula>.</p></list-item><list-item>
      <p id="d2e5221"><italic>Input configuration comparison.</italic> Three input variable combinations are evaluated at the 1, 3, and 6 h forecasting horizons to assess the relative predictive value of different information sources. Variable combination 1 comprises the five directly observed meteorological variables: AT, RH, <inline-formula><mml:math id="M208" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula>, WS, and lagged RST. Variable combination 2 extends Variable combination 1 with four ERA5-Land reanalysis variables – ST, SSR, FAL, and EVABS – selected on the basis of Spearman rank correlation with winter RST. Variable combination 3 extends Variable combination 1 with four derived features constructed exclusively from existing station observations <inline-formula><mml:math id="M209" display="inline"><mml:mrow><mml:msub><mml:mi>T</mml:mi><mml:mi mathvariant="normal">grad</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>, and RST temporal tendency features at 1, 3, and 6 h lags (<inline-formula><mml:math id="M210" display="inline"><mml:mrow><mml:mi mathvariant="normal">Δ</mml:mi><mml:msub><mml:mi mathvariant="normal">RST</mml:mi><mml:mrow><mml:mn mathvariant="normal">1</mml:mn><mml:mspace linebreak="nobreak" width="0.125em"/><mml:mi mathvariant="normal">h</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:math></inline-formula>, <inline-formula><mml:math id="M211" display="inline"><mml:mrow><mml:mi mathvariant="normal">Δ</mml:mi><mml:msub><mml:mi mathvariant="normal">RST</mml:mi><mml:mrow><mml:mn mathvariant="normal">3</mml:mn><mml:mspace linebreak="nobreak" width="0.125em"/><mml:mi mathvariant="normal">h</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:math></inline-formula>, <inline-formula><mml:math id="M212" display="inline"><mml:mrow><mml:mi mathvariant="normal">Δ</mml:mi><mml:msub><mml:mi mathvariant="normal">RST</mml:mi><mml:mrow><mml:mn mathvariant="normal">6</mml:mn><mml:mspace linebreak="nobreak" width="0.125em"/><mml:mi mathvariant="normal">h</mml:mi></mml:mrow></mml:msub></mml:mrow></mml:math></inline-formula>).</p></list-item><list-item>
      <p id="d2e5296"><italic>Feature importance and model interpretability.</italic> SHAP analysis is applied to ILES under the station-only configuration to quantify the marginal contribution of each input feature to individual predictions. Attribution values are examined stratified by temperature regime and across the diurnal cycle. The resulting feature importance patterns are qualitatively compared with expectations derived from established meteorological understanding of near-surface heat exchange processes, serving as a plausibility check on the model's learned input–output behaviour.</p></list-item><list-item>
      <p id="d2e5303"><italic>Multi-site validation.</italic> The cross-site applicability of the ILES framework is assessed at two independent stations, M9474 and M9448, using station-only inputs at the 1 h horizon. At each station, the framework is retrained independently on site-specific observations following the same protocol as the primary site. Performance rankings across five deep learning models are compared at all three stations to determine whether the predictive advantages of the ensemble are site-specific or structurally consistent across different road environments within the temperate monsoon climate zone. Uncertainty quantification provided by the BRR meta-learner is additionally assessed through a reliability diagram on the held-out test set.</p></list-item></list></p>
</sec>
</sec>
<sec id="Ch1.S4">
  <label>4</label><title>Results</title>
<sec id="Ch1.S4.SS1">
  <label>4.1</label><title>Multi-horizon prediction performance</title>
      <p id="d2e5324">Table 5 presents the prediction performance of eleven models across 1, 3, and 6 h forecasting horizons. The model hierarchy spans a naive persistence forecasting (assuming RST at time <inline-formula><mml:math id="M213" display="inline"><mml:mrow><mml:mi>t</mml:mi><mml:mo>+</mml:mo><mml:mi>h</mml:mi></mml:mrow></mml:math></inline-formula> equals RST at time <inline-formula><mml:math id="M214" display="inline"><mml:mi>t</mml:mi></mml:math></inline-formula>), a nonlinear regression baseline (NR), traditional machine learning methods (RF, Darghiasi et al., 2025; XGBoost, Kebede et al., 2024), standard deep learning architectures (GRU, LSTM, BiLSTM, CNN-LSTM, Tabrizi et al., 2021), the two proposed base learners (KNN-LSTM, BiLSTM-MHA), and the proposed ensemble ILES.</p>

<table-wrap id="T5" specific-use="star"><label>Table 5</label><caption><p id="d2e5349">Performance evaluation across 1, 3, and 6 h forecasting intervals. The last three models (KNN-LSTM, BiLSTM-MHA, ILES) represent the novel architectures proposed in this study. The preceding eight models serve as representative baselines from established model families. All models are trained and evaluated on the same dataset under the same protocol. Pairwise significance testing supporting the performance rankings reported in this table is provided in Tables S2 and S3.</p></caption><oasis:table frame="topbot"><oasis:tgroup cols="10">
     <oasis:colspec colnum="1" colname="col1" align="left"/>
     <oasis:colspec colnum="2" colname="col2" align="right"/>
     <oasis:colspec colnum="3" colname="col3" align="right"/>
     <oasis:colspec colnum="4" colname="col4" align="right" colsep="1"/>
     <oasis:colspec colnum="5" colname="col5" align="right"/>
     <oasis:colspec colnum="6" colname="col6" align="right"/>
     <oasis:colspec colnum="7" colname="col7" align="right" colsep="1"/>
     <oasis:colspec colnum="8" colname="col8" align="right"/>
     <oasis:colspec colnum="9" colname="col9" align="right"/>
     <oasis:colspec colnum="10" colname="col10" align="right"/>
     <oasis:thead>
       <oasis:row>

         <oasis:entry rowsep="1" colname="col1" morerows="1">Model</oasis:entry>

         <oasis:entry rowsep="1" namest="col2" nameend="col4" align="center" colsep="1">1 h </oasis:entry>

         <oasis:entry rowsep="1" namest="col5" nameend="col7" align="center" colsep="1">3 h </oasis:entry>

         <oasis:entry rowsep="1" namest="col8" nameend="col10" align="center">6 h </oasis:entry>

       </oasis:row>
       <oasis:row rowsep="1">

         <oasis:entry colname="col2">MAE (°C)</oasis:entry>

         <oasis:entry colname="col3">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col4">sMAPE (%)</oasis:entry>

         <oasis:entry colname="col5">MAE (°C)</oasis:entry>

         <oasis:entry colname="col6">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col7">sMAPE (%)</oasis:entry>

         <oasis:entry colname="col8">MAE (°C)</oasis:entry>

         <oasis:entry colname="col9">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col10">sMAPE (%)</oasis:entry>

       </oasis:row>
     </oasis:thead>
     <oasis:tbody>
       <oasis:row>

         <oasis:entry colname="col1">Persistence</oasis:entry>

         <oasis:entry colname="col2">0.897</oasis:entry>

         <oasis:entry colname="col3">1.295</oasis:entry>

         <oasis:entry colname="col4">35.829</oasis:entry>

         <oasis:entry colname="col5">2.497</oasis:entry>

         <oasis:entry colname="col6">3.502</oasis:entry>

         <oasis:entry colname="col7">72.193</oasis:entry>

         <oasis:entry colname="col8">4.312</oasis:entry>

         <oasis:entry colname="col9">5.788</oasis:entry>

         <oasis:entry colname="col10">101.730</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">NR</oasis:entry>

         <oasis:entry colname="col2">0.793</oasis:entry>

         <oasis:entry colname="col3">1.628</oasis:entry>

         <oasis:entry colname="col4">34.895</oasis:entry>

         <oasis:entry colname="col5">2.162</oasis:entry>

         <oasis:entry colname="col6">4.392</oasis:entry>

         <oasis:entry colname="col7">62.240</oasis:entry>

         <oasis:entry colname="col8">3.294</oasis:entry>

         <oasis:entry colname="col9">7.023</oasis:entry>

         <oasis:entry colname="col10">76.720</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">RF</oasis:entry>

         <oasis:entry colname="col2">0.524</oasis:entry>

         <oasis:entry colname="col3">0.779</oasis:entry>

         <oasis:entry colname="col4">24.777</oasis:entry>

         <oasis:entry colname="col5">1.415</oasis:entry>

         <oasis:entry colname="col6">1.975</oasis:entry>

         <oasis:entry colname="col7">51.337</oasis:entry>

         <oasis:entry colname="col8">2.259</oasis:entry>

         <oasis:entry colname="col9">3.281</oasis:entry>

         <oasis:entry colname="col10">70.928</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">XGBoost</oasis:entry>

         <oasis:entry colname="col2">0.554</oasis:entry>

         <oasis:entry colname="col3">0.871</oasis:entry>

         <oasis:entry colname="col4">26.065</oasis:entry>

         <oasis:entry colname="col5">1.401</oasis:entry>

         <oasis:entry colname="col6">1.981</oasis:entry>

         <oasis:entry colname="col7">50.967</oasis:entry>

         <oasis:entry colname="col8">2.209</oasis:entry>

         <oasis:entry colname="col9">3.248</oasis:entry>

         <oasis:entry colname="col10">70.621</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">GRU</oasis:entry>

         <oasis:entry colname="col2">0.436</oasis:entry>

         <oasis:entry colname="col3">0.630</oasis:entry>

         <oasis:entry colname="col4">22.130</oasis:entry>

         <oasis:entry colname="col5">1.345</oasis:entry>

         <oasis:entry colname="col6">1.940</oasis:entry>

         <oasis:entry colname="col7">51.071</oasis:entry>

         <oasis:entry colname="col8">2.007</oasis:entry>

         <oasis:entry colname="col9">2.792</oasis:entry>

         <oasis:entry colname="col10">67.377</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">LSTM</oasis:entry>

         <oasis:entry colname="col2">0.453</oasis:entry>

         <oasis:entry colname="col3">0.636</oasis:entry>

         <oasis:entry colname="col4">23.475</oasis:entry>

         <oasis:entry colname="col5">1.453</oasis:entry>

         <oasis:entry colname="col6">2.198</oasis:entry>

         <oasis:entry colname="col7">54.892</oasis:entry>

         <oasis:entry colname="col8">2.366</oasis:entry>

         <oasis:entry colname="col9">3.452</oasis:entry>

         <oasis:entry colname="col10">73.116</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">BiLSTM</oasis:entry>

         <oasis:entry colname="col2">0.422</oasis:entry>

         <oasis:entry colname="col3">0.572</oasis:entry>

         <oasis:entry colname="col4">22.085</oasis:entry>

         <oasis:entry colname="col5">1.416</oasis:entry>

         <oasis:entry colname="col6">1.950</oasis:entry>

         <oasis:entry colname="col7">53.451</oasis:entry>

         <oasis:entry colname="col8">2.252</oasis:entry>

         <oasis:entry colname="col9">3.092</oasis:entry>

         <oasis:entry colname="col10">72.160</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">CNN-LSTM</oasis:entry>

         <oasis:entry colname="col2">0.406</oasis:entry>

         <oasis:entry colname="col3">0.567</oasis:entry>

         <oasis:entry colname="col4">21.410</oasis:entry>

         <oasis:entry colname="col5">1.399</oasis:entry>

         <oasis:entry colname="col6">2.041</oasis:entry>

         <oasis:entry colname="col7">53.167</oasis:entry>

         <oasis:entry colname="col8">2.225</oasis:entry>

         <oasis:entry colname="col9">2.884</oasis:entry>

         <oasis:entry colname="col10">74.995</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">KNN-LSTM</oasis:entry>

         <oasis:entry colname="col2">0.389</oasis:entry>

         <oasis:entry colname="col3">0.536</oasis:entry>

         <oasis:entry colname="col4">21.348</oasis:entry>

         <oasis:entry colname="col5">1.285</oasis:entry>

         <oasis:entry colname="col6">1.926</oasis:entry>

         <oasis:entry colname="col7">47.433</oasis:entry>

         <oasis:entry colname="col8">2.170</oasis:entry>

         <oasis:entry colname="col9">2.942</oasis:entry>

         <oasis:entry colname="col10">71.853</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">BiLSTM-MHA</oasis:entry>

         <oasis:entry colname="col2">0.393</oasis:entry>

         <oasis:entry colname="col3">0.545</oasis:entry>

         <oasis:entry colname="col4">20.783</oasis:entry>

         <oasis:entry colname="col5">1.415</oasis:entry>

         <oasis:entry colname="col6">1.893</oasis:entry>

         <oasis:entry colname="col7">55.268</oasis:entry>

         <oasis:entry colname="col8">2.194</oasis:entry>

         <oasis:entry colname="col9">2.822</oasis:entry>

         <oasis:entry colname="col10">71.834</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">ILES</oasis:entry>

         <oasis:entry colname="col2">0.373</oasis:entry>

         <oasis:entry colname="col3">0.521</oasis:entry>

         <oasis:entry colname="col4">20.328</oasis:entry>

         <oasis:entry colname="col5">1.268</oasis:entry>

         <oasis:entry colname="col6">1.835</oasis:entry>

         <oasis:entry colname="col7">47.553</oasis:entry>

         <oasis:entry colname="col8">2.108</oasis:entry>

         <oasis:entry colname="col9">2.754</oasis:entry>

         <oasis:entry colname="col10">71.281</oasis:entry>

       </oasis:row>
     </oasis:tbody>
   </oasis:tgroup></oasis:table></table-wrap>

      <p id="d2e5806">At the 1 h horizon, persistence yields an MAE of 0.897 °C, RMSE of 1.295 °C, and sMAPE of 35.829 %, establishing a lower bound of predictive skill that any learned model should substantially exceed. ILES reduces MAE by 58.42 % relative to persistence and by 52.96 % relative to NR, indicating that the learned temporal representations capture substantially more information than either the most recent observation or a simple parametric reference model. Relative to the traditional machine learning baselines, RF and XGBoost, ILES reduces MAE by 28.82 % and 32.67 % respectively. Relative to the four standard deep learning architectures, GRU, LSTM, BiLSTM, and CNN-LSTM, ILES reduces MAE by 8.13 % to 17.66 %, with the smallest margin against CNN-LSTM, the strongest of the four. Relative to its own base learners, ILES reduces MAE by 4.11 % over KNN-LSTM and 5.09 % over BiLSTM-MHA.</p>
      <p id="d2e5811">At the 3 h horizon, the performance gap between persistence and all learned models widens further, consistent with the increasing value of temporal modelling at longer lead times. ILES reduces MAE by 49.22 % relative to persistence and 41.35 % relative to NR. Against the traditional machine learning and standard deep learning baselines collectively, RF, XGBoost, GRU, LSTM, BiLSTM, and CNN-LSTM, ILES reduces MAE by 5.72 % to 12.73 %. Relative to its own base learners, the margin narrows to 1.32 % over KNN-LSTM, the strongest individual model at this horizon, and 10.39 % over BiLSTM-MHA.</p>
      <p id="d2e5814">At the 6 h horizon, persistence skill degrades substantially, underscoring the practical necessity of learned models at extended lead times. ILES reduces MAE by 51.11 % relative to persistence and 36.00 % relative to NR. Among the six traditional machine learning and deep learning baselines, ILES achieves lower MAE than five, with reductions ranging from 4.57 % to 10.90 %, while GRU attains a marginally lower MAE than ILES by 5.03 %. ILES nonetheless maintains the lowest RMSE among all eleven models at this horizon. Relative to its own base learners, ILES reduces MAE by 2.86 % over KNN-LSTM and 3.92 % over BiLSTM-MHA. sMAPE values increase substantially across all models at this horizon, reflecting the prevalence of near-zero RST values in winter conditions, which inflates this metric independent of absolute forecast accuracy; MAE and RMSE remain the more informative indicators at extended horizons.</p>
      <p id="d2e5817">Across all three horizons, ILES achieves the lowest MAE and RMSE among the eleven evaluated models, with the single exception of GRU's marginally lower MAE at 6 h. sMAPE rankings are less consistent: KNN-LSTM attains a marginally lower sMAPE at 3 h, and GRU, XGBoost, and RF attain lower sMAPE at 6 h. The magnitude of improvement varies predictably with baseline category and forecasting horizon: gains over persistence and NR remain large at every horizon, ranging from 36 % to 58 %, while gains over the established machine learning and deep learning baselines narrow from roughly 8 %–33 % at 1 h to 5 %–13 % at 3 h and 5 %–11 % at 6 h, and the margin over the two base learners is smaller still, between 1 % and 10 % across horizons. This gradient is consistent with a model that captures genuine temporal structure beyond what persistence or simple regression can represent, while offering a more modest but directionally consistent refinement over architecturally related deep learning baselines and its own constituent base learners.</p>
      <p id="d2e5820">Figure 10 presents density scatter plots of predicted versus observed RST for five models across three forecasting horizons, enabling both cross-horizon and cross-model comparison of goodness-of-fit and error structure.</p>

      <fig id="F10" specific-use="star"><label>Figure 10</label><caption><p id="d2e5825">Density scatter plots of predicted versus observed winter RST across 1 h <bold>(a, d, g, j, m)</bold>, 3 h <bold>(b, e, h, k, n)</bold>, and 6 h <bold>(c, f, i, l, o)</bold> forecasting intervals. Here, colours indicate the kernel density estimation (KDE) of point concentration, with red representing high-density regions and blue representing low-density regions. Colour bar scales differ across panels to optimize the visualization of density distribution within each forecasting horizon.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f10.png"/>

        </fig>

      <p id="d2e5844">Across all models, predictive accuracy degrades systematically with forecast lead time. At the 1 h horizon, scatter points cluster tightly along the <inline-formula><mml:math id="M215" display="inline"><mml:mrow><mml:mi>y</mml:mi><mml:mo>=</mml:mo><mml:mi>x</mml:mi></mml:mrow></mml:math></inline-formula> line with high-density regions concentrated near the origin, reflecting strong model fit. At 3 and 6 h, progressive dispersion is observed: high-density regions contract, outliers increase in frequency, and regression slopes deviate increasingly from unity, indicating that all models struggle to capture rapid thermal transitions at extended horizons.</p>
      <p id="d2e5859">Across models at the 1 h horizon, LSTM achieves <inline-formula><mml:math id="M216" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup><mml:mo>=</mml:mo><mml:mn mathvariant="normal">0.9906</mml:mn></mml:mrow></mml:math></inline-formula> with a regression slope of 0.98, while BiLSTM, KNN-LSTM, and BiLSTM-MHA show progressive improvement consistent with their architectural advances. ILES attains the highest goodness-of-fit, with scatter points most tightly concentrated along the ideal line. At the 3 h horizon, the performance hierarchy is preserved: ILES achieves <inline-formula><mml:math id="M217" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup><mml:mo>=</mml:mo><mml:mn mathvariant="normal">0.9227</mml:mn></mml:mrow></mml:math></inline-formula>, outperforming LSTM, BiLSTM, KNN-LSTM, and BiLSTM-MHA. At the 6 h horizon, ILES maintains <inline-formula><mml:math id="M218" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup><mml:mo>=</mml:mo><mml:mn mathvariant="normal">0.8259</mml:mn></mml:mrow></mml:math></inline-formula>, compared to 0.7070 for LSTM, whose scatter plot exhibits the broadest dispersion and most prominent systematic underestimation at high RST values.</p>
      <p id="d2e5907">A consistent pattern across all models and horizons is that prediction errors increase with RST magnitude, as evidenced by the progressive transition from high- to low-density regions at elevated temperatures. This effect is most pronounced at the 6 h horizon, where outlier frequency increases sharply above 10 °C across all architectures, likely reflecting the greater complexity and variability of RST dynamics under daytime heating conditions, which are not directly represented in the model inputs (Qin et al., 2022). Notably, ILES exhibits the most restrained error growth with temperature, suggesting that ensemble integration partially compensates for this systematic limitation.</p>
</sec>
<sec id="Ch1.S4.SS2">
  <label>4.2</label><title>Input variable combinations comparison</title>
      <p id="d2e5918">Three input configurations are evaluated to investigate the relative predictive value of different information sources and to determine whether domain-knowledge-guided feature construction provides a practically deployable alternative to reanalysis augmentation for operational RST forecasting. <list list-type="order"><list-item>
      <p id="d2e5923">Variable combination 1 comprises AT, RH, P, WS, and historical RST, corresponding to the standard observational feature set adopted in Sect. 4.1 and serving as the reference configuration.</p></list-item><list-item>
      <p id="d2e5927">Variable combination 2 supplements Variable combination 1 with ERA5-Land-derived variables, including ST, SSR, FAL, and EVABS, representing an attempt to recover physically meaningful forcing information through grid-scale reanalysis surrogates.</p></list-item><list-item>
      <p id="d2e5931">Variable combination 3 supplements Variable combination 1 with four physically motivated features derived directly from existing station observations. The surface–air temperature difference, defined as <inline-formula><mml:math id="M219" display="inline"><mml:mrow><mml:msub><mml:mi>T</mml:mi><mml:mi mathvariant="normal">grad</mml:mi></mml:msub><mml:mo>=</mml:mo><mml:msub><mml:mi>T</mml:mi><mml:mi mathvariant="normal">RST</mml:mi></mml:msub><mml:mo>-</mml:mo><mml:msub><mml:mi>T</mml:mi><mml:mi mathvariant="normal">AT</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula>, encodes the sign and magnitude of the near-surface thermal contrast at the pavement–atmosphere interface. This quantity serves as a first-order indicator of the direction of turbulent sensible heat exchange: positive values correspond to daytime conditions under which the road surface is warmer than the overlying air, while negative values are characteristic of nocturnal radiative cooling and are commonly associated with elevated icing risk (Hermansson, 2004). The RST temporal tendency features <inline-formula><mml:math id="M220" display="inline"><mml:mrow><mml:mi mathvariant="normal">Δ</mml:mi><mml:msub><mml:mi mathvariant="normal">RST</mml:mi><mml:mi mathvariant="italic">τ</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> for <inline-formula><mml:math id="M221" display="inline"><mml:mrow><mml:mi mathvariant="italic">τ</mml:mi><mml:mo>∈</mml:mo><mml:mfenced open="{" close="}"><mml:mrow><mml:mn mathvariant="normal">1</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">3</mml:mn><mml:mo>,</mml:mo><mml:mn mathvariant="normal">6</mml:mn></mml:mrow></mml:mfenced></mml:mrow></mml:math></inline-formula> hours quantify the rate of change of road surface temperature over multiple time scales, providing explicit multi-scale descriptors of pavement thermal dynamics alongside the absolute RST sequence. <inline-formula><mml:math id="M222" display="inline"><mml:mrow><mml:msub><mml:mi>T</mml:mi><mml:mi mathvariant="normal">grad</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> provides a direct, station-measured indicator of the surface–air thermal contrast, and <inline-formula><mml:math id="M223" display="inline"><mml:mrow><mml:mi mathvariant="normal">Δ</mml:mi><mml:msub><mml:mi mathvariant="normal">RST</mml:mi><mml:mi mathvariant="italic">τ</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> provides an empirical, multi-scale descriptor of the rate of RST change; both are physically motivated by the surface energy balance framework introduced in Sect. 1, without requiring additional sensor infrastructure. Each feature combination is subsequently used as the input feature matrix <inline-formula><mml:math id="M224" display="inline"><mml:mi mathvariant="bold">X</mml:mi></mml:math></inline-formula> fed into the ILES model.</p></list-item></list></p>
      <p id="d2e6026">Performance under sub-zero conditions is assessed on the longest continuous segment of test samples satisfying RST <inline-formula><mml:math id="M225" display="inline"><mml:mo>&lt;</mml:mo></mml:math></inline-formula> 0 °C, spanning from 16:00 on 19 February 2024 to 03:00 on 22 February 2024 and comprising 60 consecutive hourly samples. This segment-based evaluation is adopted to provide a coherent assessment of model behaviour during a sustained cooling event, which constitutes the most operationally critical scenario for road icing risk assessment (Song et al., 2023; CMA, 2018).</p>
      <p id="d2e6036">As shown in Table 6, Variable combination 3 achieves the lowest MAE and RMSE across all three forecasting horizons. Its sMAPE is also lowest at the 1 and 3 h horizons. At the 1 h horizon, it attains an MAE of 0.19 °C and RMSE of 0.23 °C, representing reductions of 28.70 % and 30.10 % relative to Variable combination 1, and 42.60 % and 45.10 % relative to Variable combination 2. At 3 h, Variable combination 3 reduces MAE by 43.50 % and RMSE by 45.50 % relative to Variable combination 1, and by 48.00 % and 41.20 % relative to Variable combination 2. At 6 h, Variable combination 3 continues to yield the lowest MAE and RMSE, while Variable combination 2 remains the weakest performer across all metrics and horizons.</p>

<table-wrap id="T6" specific-use="star"><label>Table 6</label><caption><p id="d2e6043">Prediction performance of the ILES model with three input variable combinations in the subzero low temperature period.</p></caption><oasis:table frame="topbot"><oasis:tgroup cols="10">
     <oasis:colspec colnum="1" colname="col1" align="left"/>
     <oasis:colspec colnum="2" colname="col2" align="right"/>
     <oasis:colspec colnum="3" colname="col3" align="right"/>
     <oasis:colspec colnum="4" colname="col4" align="right" colsep="1"/>
     <oasis:colspec colnum="5" colname="col5" align="right"/>
     <oasis:colspec colnum="6" colname="col6" align="right"/>
     <oasis:colspec colnum="7" colname="col7" align="right" colsep="1"/>
     <oasis:colspec colnum="8" colname="col8" align="right"/>
     <oasis:colspec colnum="9" colname="col9" align="right"/>
     <oasis:colspec colnum="10" colname="col10" align="right"/>
     <oasis:thead>
       <oasis:row>

         <oasis:entry rowsep="1" colname="col1" morerows="1">Input variable combinations</oasis:entry>

         <oasis:entry rowsep="1" namest="col2" nameend="col4" align="center" colsep="1">1 h </oasis:entry>

         <oasis:entry rowsep="1" namest="col5" nameend="col7" align="center" colsep="1">3 h </oasis:entry>

         <oasis:entry rowsep="1" namest="col8" nameend="col10" align="center">6 h </oasis:entry>

       </oasis:row>
       <oasis:row rowsep="1">

         <oasis:entry colname="col2">MAE (°C)</oasis:entry>

         <oasis:entry colname="col3">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col4">sMAPE (%)</oasis:entry>

         <oasis:entry colname="col5">MAE (°C)</oasis:entry>

         <oasis:entry colname="col6">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col7">sMAPE (%)</oasis:entry>

         <oasis:entry colname="col8">MAE (°C)</oasis:entry>

         <oasis:entry colname="col9">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col10">sMAPE (%)</oasis:entry>

       </oasis:row>
     </oasis:thead>
     <oasis:tbody>
       <oasis:row>

         <oasis:entry colname="col1">1</oasis:entry>

         <oasis:entry colname="col2">0.272</oasis:entry>

         <oasis:entry colname="col3">0.329</oasis:entry>

         <oasis:entry colname="col4">15.545</oasis:entry>

         <oasis:entry colname="col5">1.860</oasis:entry>

         <oasis:entry colname="col6">2.410</oasis:entry>

         <oasis:entry colname="col7">90.851</oasis:entry>

         <oasis:entry colname="col8">2.063</oasis:entry>

         <oasis:entry colname="col9">2.623</oasis:entry>

         <oasis:entry colname="col10">91.885</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">2</oasis:entry>

         <oasis:entry colname="col2">0.338</oasis:entry>

         <oasis:entry colname="col3">0.419</oasis:entry>

         <oasis:entry colname="col4">17.423</oasis:entry>

         <oasis:entry colname="col5">2.020</oasis:entry>

         <oasis:entry colname="col6">2.230</oasis:entry>

         <oasis:entry colname="col7">93.515</oasis:entry>

         <oasis:entry colname="col8">1.963</oasis:entry>

         <oasis:entry colname="col9">2.489</oasis:entry>

         <oasis:entry colname="col10">97.241</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">3</oasis:entry>

         <oasis:entry colname="col2">0.194</oasis:entry>

         <oasis:entry colname="col3">0.230</oasis:entry>

         <oasis:entry colname="col4">8.485</oasis:entry>

         <oasis:entry colname="col5">1.050</oasis:entry>

         <oasis:entry colname="col6">1.312</oasis:entry>

         <oasis:entry colname="col7">63.826</oasis:entry>

         <oasis:entry colname="col8">1.726</oasis:entry>

         <oasis:entry colname="col9">1.912</oasis:entry>

         <oasis:entry colname="col10">100.355</oasis:entry>

       </oasis:row>
     </oasis:tbody>
   </oasis:tgroup></oasis:table></table-wrap>

      <p id="d2e6227">Figure 11 corroborates this hierarchy visually. At the 1 h horizon, Variable combination 3 tracks the observed cooling descent from 20 to 21 February most faithfully, while combinations 1 and 2 exhibit systematic underestimation of the cooling rate. This behaviour is consistent with the role of <inline-formula><mml:math id="M226" display="inline"><mml:mrow><mml:msub><mml:mi>T</mml:mi><mml:mi mathvariant="normal">grad</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> and <inline-formula><mml:math id="M227" display="inline"><mml:mrow><mml:mi mathvariant="normal">Δ</mml:mi><mml:msub><mml:mi mathvariant="normal">RST</mml:mi><mml:mi mathvariant="italic">τ</mml:mi></mml:msub></mml:mrow></mml:math></inline-formula> in supplying explicit thermal contrast and rate-of-change information during periods of sustained temperature decline, when these signals are most informative. At 3 and 6 h horizons, all configurations show progressive amplitude attenuation; Variable combination 2 exhibits the most pronounced phase lag and amplitude underestimation, while Variable combination 3 retains the closest agreement with observed thermal minima.</p>

      <fig id="F11" specific-use="star"><label>Figure 11</label><caption><p id="d2e6256">Winter RST prediction of ILES model with three input variable combinations in subzero low temperature period across 1 h <bold>(a)</bold>, 3 h <bold>(b)</bold>, and 6 h <bold>(c)</bold> forecasting intervals. This period is the longest continuous section of the test sample with RST <inline-formula><mml:math id="M228" display="inline"><mml:mo>&lt;</mml:mo></mml:math></inline-formula> 0  °C.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f11.png"/>

        </fig>

      <p id="d2e6281">The degraded performance of Variable combination 2 is attributable to two mechanisms. First, ERA5-Land provides area-averaged estimates at approximately 9 to 11 km resolution, whereas RST is governed by point-scale surface conditions; this spatial representativeness mismatch introduces estimation errors that reduce rather than enhance predictive content, a well-documented limitation of gridded reanalysis products in surface applications (Muñoz-Sabater et al., 2021). Second, the historical RST sequence included in all three configurations already implicitly encodes the net effect of prior meteorological forcing at the pavement surface, rendering the reanalysis-derived surrogates largely redundant while simultaneously increasing input dimensionality (Reichstein et al., 2019). These conclusions are specific to the instrumented, data-rich setting examined here; in station-sparse regions, reanalysis variables may retain supplementary predictive value provided that appropriate bias correction is applied to mitigate the identified scale mismatch.</p>
      <p id="d2e6284">To assess whether the performance hierarchy established under sub-zero conditions reflects a general advantage of Variable combination 3 rather than a regime-specific result, two three-day periods are further examined: a stable clear-sky period (25 to 27 January 2024, characterised by mean relative humidity below 50 %, zero precipitation, and wind speeds below 2 m s<sup>−1</sup>) and an overcast rainy period (23 to 25 February 2024, characterised by continuous precipitation and relative humidity exceeding 85 %). These two periods represent contrasting RST variability regimes, spanning periodically dominated and stochastically driven conditions, with direct relevance to operational winter road maintenance (Darghiasi et al., 2023).</p>
      <p id="d2e6300">Under overcast and rainy conditions (Fig. 12a to c), RST variability becomes markedly more subdued and irregular. Variable combination 2 exhibits the most severe performance degradation across all horizons, with persistent positive bias and substantial loss of observed thermal structure at 6 h, attributable to the introduction of spatially mismatched reanalysis information under precipitation-dominated conditions. Variable combination 3 consistently maintains the closest correspondence with observations, with its relative advantage most pronounced at 3 and 6 h.</p>

      <fig id="F12" specific-use="star"><label>Figure 12</label><caption><p id="d2e6305">Winter RST prediction of ILES model with three input variables in overcast and rainy and clear synoptic conditions across 1 h <bold>(a, d)</bold>, 3 h <bold>(b, e)</bold>, and 6 h <bold>(c, f)</bold> forecasting intervals.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f12.png"/>

        </fig>

      <p id="d2e6323">Under clear-sky conditions (Fig. 12d to f), RST exhibits pronounced diurnal oscillations with peak-to-trough amplitudes of approximately 17 °C. At the 1 h horizon, all three configurations achieve comparable accuracy. At 3 and 6 h, Variable combination 2 shows progressive amplitude underestimation, while Variable combination 3 maintains the closest alignment with the observed diurnal cycle, consistent with the high predictive content of multi-scale tendency features under regular periodic forcing.</p>
      <p id="d2e6326">Taken together, the results across sub-zero, clear-sky, and overcast regimes consistently favour Variable combination 3, confirming that physics-motivated feature engineering provides the most effective input strategy for station-based winter RST prediction at instrumented sites. Accordingly, Variable combination 1 is retained as the standard input configuration for Sect. 4.3 and 4.4. Since all ten benchmark models in Sect. 4.1 were evaluated under this configuration, retaining it as the reference ensures that SHAP-based attribution and cross-site performance rankings remain directly comparable to the established benchmarks without confounding from input differences. It should be noted that in operational settings where physics-motivated feature construction is feasible, variable combination 3 is recommended as the preferred input strategy given its consistently lowest MAE and RMSE across all forecasting horizons and meteorological regimes.</p>
</sec>
<sec id="Ch1.S4.SS3">
  <label>4.3</label><title>Feature importance and model interpretability</title>
      <p id="d2e6337">To examine the degree to which the learned input–output relationships of the ILES model are consistent with established meteorological understanding of RST dynamics, SHAP analysis is applied under the station-only input configuration (Variable combination 1). Since the BRR meta-learner operates on base learner predictions rather than raw meteorological features, SHAP values cannot be applied directly to the ensemble output to attribute importance to the original inputs. Instead, KernelSHAP (Lundberg and Lee, 2017) is applied independently to each base learner (KNN-LSTM and BiLSTM-MHA) with respect to the original input feature set, treating each base learner as a black-box function mapping raw features to RST predictions. The resulting SHAP values from the two base learners are aggregated as a weighted average, with weights proportional to the BRR meta-learner coefficients, to produce a single set of feature attributions representative of the full ensemble. This procedure provides a model-agnostic and theoretically principled attribution of ensemble predictions to the original meteorological inputs, grounded in cooperative game theory (Lundberg and Lee, 2017; Joo et al., 2023). Attribution values are examined stratified by temperature regime and across the diurnal cycle, enabling a systematic assessment of whether the learned importance structure is qualitatively consistent with the dominant meteorological drivers identified by surface energy balance theory.</p>
      <p id="d2e6340">The SHAP analysis presented here is restricted to the four instantaneous meteorological input variables (AT, RH, WS, <inline-formula><mml:math id="M230" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula>). The reported SHAP importances therefore characterise the marginal contributions of the contemporaneous meteorological state, not the full input space. AT is the dominant predictor (Fig. 13a), with a mean absolute SHAP value of 0.696, accounting for 55.4 % of total feature importance. RH, WS, and <inline-formula><mml:math id="M231" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula> follow in descending order, contributing 16.9 %, 14.3 %, and 13.4 % respectively. This ranking is consistent with the recognized role of AT as the primary driver of RST through near-surface sensible heat exchange, and the secondary modulating roles of RH, WS, and <inline-formula><mml:math id="M232" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula> through evaporative, convective, and latent heat processes (Chen et al., 2019; Gui et al., 2007; Feng and Feng, 2012). The beeswarm plot (Fig. 13b) confirms the expected directionality: high AT values are associated with positive SHAP contributions across the full temperature range, while elevated RH and <inline-formula><mml:math id="M233" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula> values are predominantly associated with negative contributions, consistent with their cooling influence under high-humidity and wet-surface conditions. The narrow spread of <inline-formula><mml:math id="M234" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula> points reflects its episodic character, with predictive influence concentrated in a small fraction of precipitation hours.</p>

      <fig id="F13" specific-use="star"><label>Figure 13</label><caption><p id="d2e6380">Mean absolute SHAP values <bold>(a)</bold> and beeswarm plots of SHAP value distributions <bold>(b)</bold> for the ILES model.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f13.png"/>

        </fig>

      <p id="d2e6396">Figure 14a presents the diurnal evolution of mean absolute SHAP values. AT maintains the highest importance at all hours, with moderately elevated values during the solar heating window (09:00 to 16:00) and the nocturnal cooling window (22:00 to 06:00), broadly consistent with the larger surface-air thermal contrast expected during these periods. RH, WS, and <inline-formula><mml:math id="M235" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula> display comparatively stable diurnal profiles without systematic hour-to-hour variation. Figure 14b presents importance values stratified by RST regime. AT importance is notably elevated at thermal extremes, with mean absolute SHAP values of 1.603 at RST below <inline-formula><mml:math id="M236" display="inline"><mml:mo>-</mml:mo></mml:math></inline-formula>5 °C and 1.380 at RST above 15 °C, compared to 0.224 in the 0 to 5 °C range, consistent with the larger surface-air temperature differences expected under extreme thermal conditions. In the near-neutral regime (0 to 5 °C), feature importances are compressed across all variables, indicating reduced dominance of any single predictor. In the above-15 °C regime, <inline-formula><mml:math id="M237" display="inline"><mml:mi>P</mml:mi></mml:math></inline-formula> rises to second rank with a mean absolute SHAP value of 0.419, though the small sample size in this stratum (<inline-formula><mml:math id="M238" display="inline"><mml:mrow><mml:mi>n</mml:mi><mml:mo>=</mml:mo><mml:mn mathvariant="normal">56</mml:mn></mml:mrow></mml:math></inline-formula>) warrants caution in interpretation.</p>

      <fig id="F14" specific-use="star"><label>Figure 14</label><caption><p id="d2e6434">Diurnal variation of mean absolute SHAP values across the 24 h cycle <bold>(a)</bold>. Shaded regions indicate the solar heating window (09:00–16:00) and nocturnal cooling window (22:00–06:00). Mean absolute SHAP values stratified by RST regime <bold>(b)</bold>. Numbers in parentheses indicate sample counts per regime.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f14.png"/>

        </fig>

      <p id="d2e6449">Collectively, the SHAP results indicate that the ILES model has learned importance rankings that are qualitatively consistent with the primary meteorological drivers of RST identified by surface energy balance theory, across both precipitation states and temperature regimes. It should be noted that SHAP analysis characterises learned statistical associations and does not establish causal correspondence with physical processes; the consistency identified here is a necessary but not sufficient condition for physical plausibility.</p>
</sec>
<sec id="Ch1.S4.SS4">
  <label>4.4</label><title>Multi-site validation</title>
      <p id="d2e6460">To assess the cross-site applicability of the proposed ILES framework, validation was conducted at two independent stations: M9474 (Xiaohuangshan, 32.04° N, 119.86° E) and M9448 (Huai'an Airport, 33.75° N, 119.17° E), both located on the central Jiangsu plain. M9474 is situated on a cross-Yangtze River bridge, approximately 363 km southeast of the primary station; M9448 is located near Huai'an Airport, approximately 206 km southeast of M9393, on flat open terrain. As with the primary station M9393, both M9474 and M9448 are situated in relatively unobstructed settings without significant roadside screening structures; this validation therefore assesses transferability of the framework's architecture and training protocol across stations sharing broadly similar exposure geometries, rather than across the specific shading or sky-view conditions encountered at screened road segments. The input variables at both stations comprise site-specific air temperature, wind speed, relative humidity, precipitation, and historical RST. Consistent with the imputation procedure described in Sect. 3.1.1, no short-gap event in the test period has a cross-hour interpolation anchor at either validation station, and long-gap climatological fills – computed using the same station-specific, training-period-only procedure – account for 0.66 % and 0.13 %  of 5 min test-period samples at M9474 and M9448, respectively. For M9474, the winter of 2024 served as the test set. For M9448, three winter seasons (January–February 2017, December 2017–February 2018, and December 2018–February 2019) were used for training, with December 2019–February 2020 as the test set. At each site, the ILES framework was retrained independently on site-specific observations following the same experimental protocol as described in Sect. 3.3.2.</p>
      <p id="d2e6463">Table 7 presents the 1 h prediction results at all three stations. ILES consistently outperforms the standard LSTM model at all three sites, with MAE reductions of 17.66 %, 5.68 %, and 35.65 % at M9393, M9474, and M9448, respectively. The consistent performance advantage of ILES across stations with distinct surface conditions confirms the structural robustness of the proposed approach within the temperate monsoon climate zone.</p>

<table-wrap id="T7" specific-use="star"><label>Table 7</label><caption><p id="d2e6469">Comparison of MAE, RMSE, and sMAPE for 1 h winter RST prediction at M9393, M9474 and M9448 sites using five deep learning models.</p></caption><oasis:table frame="topbot"><oasis:tgroup cols="10">
     <oasis:colspec colnum="1" colname="col1" align="left"/>
     <oasis:colspec colnum="2" colname="col2" align="center"/>
     <oasis:colspec colnum="3" colname="col3" align="center"/>
     <oasis:colspec colnum="4" colname="col4" align="center" colsep="1"/>
     <oasis:colspec colnum="5" colname="col5" align="center"/>
     <oasis:colspec colnum="6" colname="col6" align="center"/>
     <oasis:colspec colnum="7" colname="col7" align="center" colsep="1"/>
     <oasis:colspec colnum="8" colname="col8" align="center"/>
     <oasis:colspec colnum="9" colname="col9" align="center"/>
     <oasis:colspec colnum="10" colname="col10" align="center"/>
     <oasis:thead>
       <oasis:row>

         <oasis:entry rowsep="1" colname="col1" morerows="1">Model</oasis:entry>

         <oasis:entry rowsep="1" namest="col2" nameend="col4" colsep="1">M9393 </oasis:entry>

         <oasis:entry rowsep="1" namest="col5" nameend="col7" colsep="1">M9474 </oasis:entry>

         <oasis:entry rowsep="1" namest="col8" nameend="col10">M9448 </oasis:entry>

       </oasis:row>
       <oasis:row rowsep="1">

         <oasis:entry colname="col2">MAE (°C)</oasis:entry>

         <oasis:entry colname="col3">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col4">sMAPE (%)</oasis:entry>

         <oasis:entry colname="col5">MAE (°C)</oasis:entry>

         <oasis:entry colname="col6">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col7">sMAPE (%)</oasis:entry>

         <oasis:entry colname="col8">MAE (°C)</oasis:entry>

         <oasis:entry colname="col9">RMSE (°C)</oasis:entry>

         <oasis:entry colname="col10">sMAPE (%)</oasis:entry>

       </oasis:row>
     </oasis:thead>
     <oasis:tbody>
       <oasis:row>

         <oasis:entry colname="col1">LSTM</oasis:entry>

         <oasis:entry colname="col2">0.453</oasis:entry>

         <oasis:entry colname="col3">0.636</oasis:entry>

         <oasis:entry colname="col4">23.475</oasis:entry>

         <oasis:entry colname="col5">0.229</oasis:entry>

         <oasis:entry colname="col6">0.330</oasis:entry>

         <oasis:entry colname="col7">5.109</oasis:entry>

         <oasis:entry colname="col8">0.620</oasis:entry>

         <oasis:entry colname="col9">0.793</oasis:entry>

         <oasis:entry colname="col10">17.639</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">BiLSTM</oasis:entry>

         <oasis:entry colname="col2">0.422</oasis:entry>

         <oasis:entry colname="col3">0.572</oasis:entry>

         <oasis:entry colname="col4">22.085</oasis:entry>

         <oasis:entry colname="col5">0.228</oasis:entry>

         <oasis:entry colname="col6">0.358</oasis:entry>

         <oasis:entry colname="col7">4.182</oasis:entry>

         <oasis:entry colname="col8">0.433</oasis:entry>

         <oasis:entry colname="col9">0.629</oasis:entry>

         <oasis:entry colname="col10">11.801</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">KNN-LSTM</oasis:entry>

         <oasis:entry colname="col2">0.389</oasis:entry>

         <oasis:entry colname="col3">0.536</oasis:entry>

         <oasis:entry colname="col4">21.348</oasis:entry>

         <oasis:entry colname="col5">0.218</oasis:entry>

         <oasis:entry colname="col6">0.331</oasis:entry>

         <oasis:entry colname="col7">4.204</oasis:entry>

         <oasis:entry colname="col8">0.420</oasis:entry>

         <oasis:entry colname="col9">0.629</oasis:entry>

         <oasis:entry colname="col10">11.428</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">BiLSTM-MHA</oasis:entry>

         <oasis:entry colname="col2">0.393</oasis:entry>

         <oasis:entry colname="col3">0.545</oasis:entry>

         <oasis:entry colname="col4">20.783</oasis:entry>

         <oasis:entry colname="col5">0.224</oasis:entry>

         <oasis:entry colname="col6">0.330</oasis:entry>

         <oasis:entry colname="col7">4.317</oasis:entry>

         <oasis:entry colname="col8">0.431</oasis:entry>

         <oasis:entry colname="col9">0.617</oasis:entry>

         <oasis:entry colname="col10">12.487</oasis:entry>

       </oasis:row>
       <oasis:row>

         <oasis:entry colname="col1">ILES</oasis:entry>

         <oasis:entry colname="col2">0.373</oasis:entry>

         <oasis:entry colname="col3">0.521</oasis:entry>

         <oasis:entry colname="col4">20.328</oasis:entry>

         <oasis:entry colname="col5">0.216</oasis:entry>

         <oasis:entry colname="col6">0.329</oasis:entry>

         <oasis:entry colname="col7">4.487</oasis:entry>

         <oasis:entry colname="col8">0.399</oasis:entry>

         <oasis:entry colname="col9">0.599</oasis:entry>

         <oasis:entry colname="col10">10.883</oasis:entry>

       </oasis:row>
     </oasis:tbody>
   </oasis:tgroup></oasis:table></table-wrap>

      <p id="d2e6722">The BRR meta-learner yields closed-form probabilistic forecasts through its posterior predictive distribution (Eq. 20). Figure 15 presents reliability diagrams evaluated on the held-out test sets at all three stations, in which the empirical coverage curves lie consistently above the diagonal of perfect calibration across all nominal confidence levels, indicating systematic conservative over-coverage at each site. At M9393, the 68 %, 90 %, and 95 % prediction intervals achieve empirical coverages of 88.1 %, 96.8 %, and 98.1 %, with mean interval widths of 1.523, 2.519, and 3.000 °C, respectively. At M9474, the corresponding coverages are 88.1 %, 95.6 %, and 97.5 %, with mean widths of 0.967, 1.600, and 1.906 °C. At M9448, coverages of 88.1 %, 95.6 %, and 97.5 % are achieved with mean widths of 2.216, 3.666, and 4.368 °C. The notably narrower intervals at M9474 relative to M9393 reflect the lower RST variability characteristic of the bridge-mounted station, while the consistent coverage hierarchy across all three sites confirms that the conservative calibration pattern is a structural property of the BRR posterior predictive decomposition (Eq. 20) rather than a site-specific artefact.</p>

      <fig id="F15" specific-use="star"><label>Figure 15</label><caption><p id="d2e6727">Reliability diagrams of the BRR meta-learner for 1 h winter RST prediction at M9393, M9474, and M9448.</p></caption>
          <graphic xlink:href="https://gmd.copernicus.org/articles/19/9035/2026/gmd-19-9035-2026-f15.png"/>

        </fig>

</sec>
</sec>
<sec id="Ch1.S5" sec-type="conclusions">
  <label>5</label><title>Conclusion</title>
      <p id="d2e6746">This study proposed the Improved LSTMs Ensemble with Stacking (ILES) framework for winter road surface temperature prediction, integrating two base learners with partially complementary predictive characteristics within a stacking ensemble trained via three-fold temporal cross-validation aligned to complete winter seasons. The framework was evaluated across four analytical perspectives using four consecutive winters of observations from an operational road meteorological station in Jiangsu, China.</p>
      <p id="d2e6749">On predictive performance, ILES generally achieved the lowest prediction errors among all eleven evaluated models at each of the 1, 3, and 6 h forecasting horizons, with MAE ranging from 0.373 °C at 1 h to 2.108 °C at 6 h and <inline-formula><mml:math id="M239" display="inline"><mml:mrow><mml:msup><mml:mi>R</mml:mi><mml:mn mathvariant="normal">2</mml:mn></mml:msup></mml:mrow></mml:math></inline-formula> reaching 0.993, 0.923, and 0.826, respectively. Against the six established machine learning and deep learning baselines – RF, XGBoost, GRU, LSTM, BiLSTM, and CNN-LSTM – ILES reduced MAE by 8.13 %–32.67 % and RMSE by 8.11 %–40.18 % at 1 h, and cut MAE by 5.72 %–12.73 % and RMSE by 5.41 %–16.51 % at 3 h; at 6 h, ILES achieved lower MAE and RMSE than five of the six established baselines, the exception being GRU, though ILES retained the lowest RMSE across all eleven models. sMAPE rankings were not uniformly optimal: at the 6 h horizon, GRU, XGBoost, and RF achieved lower sMAPE than ILES, a pattern attributable to the disproportionate sensitivity of sMAPE to near-zero RST values at extended horizons.</p>
      <p id="d2e6763">On input configuration, physics-motivated feature engineering incorporating the surface–air temperature difference and multi-scale RST temporal tendency features achieved the lowest MAE and RMSE relative to both the station-only observational baseline and ERA5-Land reanalysis augmentation across all forecasting horizons and meteorological regimes examined, including sub-zero, overcast, and clear-sky conditions. These findings demonstrate that meaningful predictive gains at instrumented sites can be achieved through domain-knowledge-guided feature construction alone, without additional sensor infrastructure or external reanalysis data products, and establish Variable combination 3 as the recommended input strategy for operational deployment at monitored road segments. On model interpretability, stratified SHAP analysis confirmed that the learned feature importance rankings are qualitatively consistent with the dominant meteorological drivers identified by surface energy balance theory across temperature regimes and the diurnal cycle, with air temperature as the primary predictor and regime-dependent rank shifts among secondary predictors.</p>
      <p id="d2e6766">On multi-site applicability, ILES maintained its predictive superiority over the LSTM baseline at both independent validation stations within the same temperate monsoon climate zone of Jiangsu Province, with MAE reductions of 17.66 %, 5.68 %, and 35.65 % at M9393, M9474, and M9448, respectively. The BRR meta-learner additionally yielded conservative probabilistic forecasts with empirical prediction interval coverages consistently exceeding nominal confidence levels across all three sites, a calibration pattern confirmed to be a structural property of the BRR posterior predictive decomposition rather than a site-specific artefact.</p>
      <p id="d2e6770">Several limitations motivate future work. The framework was evaluated under a temperate monsoon climate, and its transferability to subarctic, alpine, or maritime regimes remains to be assessed. The absence of direct shortwave radiation measurements and local exposure descriptors – such as sky-view factor, roadside screening geometry, and road orientation – means that the model is best characterised as a station-specific time-series predictor at a fixed exposure. The learned diurnal pattern in the RST input sequence is conditioned on the specific radiation environment at each monitored cross-section and should not be extrapolated to nearby road segments with substantially different shading or orientation without site-specific retraining or the inclusion of explicit radiation inputs. Integration of numerical weather prediction outputs to extend forecast horizons beyond 6 h, and the application of physics-informed architectural constraints for data-sparse settings, represent productive directions for future development.</p>
</sec>

      
      </body>
    <back><notes notes-type="codedataavailability"><title>Code and data availability</title>

      <p id="d2e6777">The codes for conducting the analyses can be downloaded from <ext-link xlink:href="https://doi.org/10.5281/zenodo.22020890" ext-link-type="DOI">10.5281/zenodo.22020890</ext-link> (Li, 2026). The ERA5-Land reanalysis data are available from the Copernicus Climate Change Service (C3S) Climate Data Store at <ext-link xlink:href="https://doi.org/10.24381/cds.e2161bac" ext-link-type="DOI">10.24381/cds.e2161bac</ext-link> (Muñoz Sabater, 2019). All data used in this study are publicly available.</p>
  </notes><app-group>
        <supplementary-material position="anchor"><p id="d2e6786">The supplement related to this article is available online at <inline-supplementary-material xlink:href="https://doi.org/10.5194/gmd-19-9035-2026-supplement" xlink:title="pdf">https://doi.org/10.5194/gmd-19-9035-2026-supplement</inline-supplementary-material>.</p></supplementary-material>
        </app-group><notes notes-type="authorcontribution"><title>Author contributions</title>

      <p id="d2e6795">WTL, LYZ and XHW designed the study. WTL developed the model and wrote the paper. XHW, YHG, YHG, KC, WQH and WQH analysed the data. All authors contribute to writing the paper.</p>
  </notes><notes notes-type="competinginterests"><title>Competing interests</title>

      <p id="d2e6801">The contact author has declared that none of the authors has any competing interests.</p>
  </notes><notes notes-type="disclaimer"><title>Disclaimer</title>

      <p id="d2e6807">Publisher's note: Copernicus Publications remains neutral with regard to jurisdictional claims made in the text, published maps, institutional affiliations, or any other geographical representation in this paper. The authors bear the ultimate responsibility for providing appropriate place names. Views expressed in the text are those of the authors and do not necessarily reflect the views of the publisher.</p>
  </notes><notes notes-type="financialsupport"><title>Financial support</title>

      <p id="d2e6813">This research has been supported by the National Natural Science Foundation of China (grant nos. U24A20606, 42575210, 41975087, and 42075068), the State Key Laboratory of Severe Weather, China (grant no. 2023LASW-B25), the Yunnan Provincial Department of Transportation, China (grant no. 2023-152), and the Nanjing University of Information Science and Technology, China (grant no. 2025h522).</p>
  </notes><notes notes-type="reviewstatement"><title>Review statement</title>

      <p id="d2e6819">This paper was edited by Patricia Lawston-Parker and reviewed by Mahya Hashemi and one anonymous referee.</p>
  </notes><ref-list>
    <title>References</title>

      <ref id="bib1.bib1"><label>1</label><mixed-citation>Abo-Hashema, M. A.: Modeling pavement temperature prediction using artificial neural networks, in: Airfield and highway pavement 2013: Sustainable and efficient pavements, American Society of Civil Engineers, Reston, VA, USA, 490–505,  <ext-link xlink:href="https://doi.org/10.1061/9780784413005.039" ext-link-type="DOI">10.1061/9780784413005.039</ext-link>, 2013.</mixed-citation></ref>
      <ref id="bib1.bib2"><label>2</label><mixed-citation>Adwan, I., Milad, A., Memon, Z. A., Widyatmoko, I., Zanuri, N. A., Memon, N. A., and Yusoff, N. I. M.:  Asphalt pavement temperature prediction models: A review, Appl. Sci., 11, 3794, <ext-link xlink:href="https://doi.org/10.3390/app11093794" ext-link-type="DOI">10.3390/app11093794</ext-link>, 2021.</mixed-citation></ref>
      <ref id="bib1.bib3"><label>3</label><mixed-citation>Asefzadeh, A., Hashemian, L., and Bayat, A.: Development of statistical temperature prediction models for a test road in Edmonton, Alberta, Canada, Int. J. Pavement Res. Technol., 10, 369–382, <ext-link xlink:href="https://doi.org/10.1016/j.ijprt.2017.05.003" ext-link-type="DOI">10.1016/j.ijprt.2017.05.003</ext-link>, 2017.</mixed-citation></ref>
      <ref id="bib1.bib4"><label>4</label><mixed-citation>Athukorallage, B., Senadheera, S., and James, D.: Temporal and spatial temperature predictions for flexible pavement layers using numerical thermal analysis and verified with large datasets, Case Stud. Constr. Mater., 18, e02008, <ext-link xlink:href="https://doi.org/10.1016/j.cscm.2023.e02008" ext-link-type="DOI">10.1016/j.cscm.2023.e02008</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib5"><label>5</label><mixed-citation>Ayasrah, U. B., Tashman, L., AlOmari, A., and Asi, I.: Development of a temperature prediction model for flexible pavement structures, Case Stud. Constr. Mater., 18, e01697, <ext-link xlink:href="https://doi.org/10.1016/j.cscm.2022.e01697" ext-link-type="DOI">10.1016/j.cscm.2022.e01697</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib6"><label>6</label><mixed-citation>Ba, J. L., Kiros, J. R., and Hinton, G. E.: Layer normalization, arXiv [preprint], <ext-link xlink:href="https://doi.org/10.48550/arXiv.1607.06450" ext-link-type="DOI">10.48550/arXiv.1607.06450</ext-link>, 2016.</mixed-citation></ref>
      <ref id="bib1.bib7"><label>7</label><mixed-citation>Bai, S., Yang, W., Zhang, M., Liu, D., Li, W., and Zhou, L.: Attention-based BiLSTM model for pavement temperature prediction of asphalt pavement in winter, Atmosphere, 13, 1524, <ext-link xlink:href="https://doi.org/10.3390/atmos13091524" ext-link-type="DOI">10.3390/atmos13091524</ext-link>, 2022.</mixed-citation></ref>
      <ref id="bib1.bib8"><label>8</label><mixed-citation>Ben Taieb, S., Bontempi, G., Atiya, A. F., and Sorjamaa, A.: A review and comparison of strategies for multi-step ahead time series forecasting based on the NN5 forecasting competition, Expert Syst. Appl., 39, 7067–7083, <ext-link xlink:href="https://doi.org/10.1016/j.eswa.2012.01.039" ext-link-type="DOI">10.1016/j.eswa.2012.01.039</ext-link>, 2012.</mixed-citation></ref>
      <ref id="bib1.bib9"><label>9</label><mixed-citation>Bishop, C. M., Nasrabadi, N. M.: Pattern recognition and machine learning, New York, Springer, <ext-link xlink:href="https://doi.org/10.1080/15228053.2019.1632410" ext-link-type="DOI">10.1080/15228053.2019.1632410</ext-link>, 2006.</mixed-citation></ref>
      <ref id="bib1.bib10"><label>10</label><mixed-citation>Breiman, L.: Bagging predictors, Mach. Learn., 24, 123–140, <ext-link xlink:href="https://doi.org/10.1007/BF00058655" ext-link-type="DOI">10.1007/BF00058655</ext-link>, 1996.</mixed-citation></ref>
      <ref id="bib1.bib11"><label>11</label><mixed-citation>Chen, J., Wang, H., and Xie, P.: Pavement temperature prediction: Theoretical models and critical affecting factors, Appl. Therm. Eng., 158, 113755, <ext-link xlink:href="https://doi.org/10.1016/j.applthermaleng.2019.113755" ext-link-type="DOI">10.1016/j.applthermaleng.2019.113755</ext-link>, 2019.</mixed-citation></ref>
      <ref id="bib1.bib12"><label>12</label><mixed-citation>Cheng, H., Liu, J., Sun, L., and Liu, L.: Critical position of fatigue damage within asphalt pavement considering temperature and strain distribution, Int. J. Pavement Eng., 22, 1773–1784, <ext-link xlink:href="https://doi.org/10.1080/10298436.2020.1724288" ext-link-type="DOI">10.1080/10298436.2020.1724288</ext-link>, 2021.</mixed-citation></ref>
      <ref id="bib1.bib13"><label>13</label><mixed-citation> China Meteorological Administration: Grade of highway traffic high-impact weather warning, QX/T 414-2018, China Meteorological Press, Beijing, China, 2018.</mixed-citation></ref>
      <ref id="bib1.bib14"><label>14</label><mixed-citation>Crevier, L.-P. and Delage, Y.: METRo: A new model for road-condition forecasting in Canada, J. Appl. Meteorol., 40, 2026–2037, <ext-link xlink:href="https://doi.org/10.1175/1520-0450(2001)040&lt;2026:MANMFR&gt;2.0.CO;2" ext-link-type="DOI">10.1175/1520-0450(2001)040&lt;2026:MANMFR&gt;2.0.CO;2</ext-link>, 2001.</mixed-citation></ref>
      <ref id="bib1.bib15"><label>15</label><mixed-citation>Dai, B., Yang, W., Ji, X., and Zhou, L.: An ensemble deep learning model for short-term road surface temperature prediction, J. Transp. Eng. B-Pavements, 149, 04022067, <ext-link xlink:href="https://doi.org/10.1061/JPEODX.PVENG-1192" ext-link-type="DOI">10.1061/JPEODX.PVENG-1192</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib16"><label>16</label><mixed-citation>Darghiasi, P., Baral, A., Mattingly, S., and Shahandashti, M.: Estimation of road surface temperature using NOAA gridded forecast weather data for snowplow operations management, J. Cold Reg. Eng., 37, 04023018, <ext-link xlink:href="https://doi.org/10.1061/JCRGEI.CRENG-691" ext-link-type="DOI">10.1061/JCRGEI.CRENG-691</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib17"><label>17</label><mixed-citation>Darghiasi, P., Zamanian, M., and Shahandashti, M.: Enhancing Winter Maintenance Decision Making through Deep Learning-Based Road Surface Temperature Estimation, in: Construction Res. Congr. 2024, ASCE, 701–711, <ext-link xlink:href="https://doi.org/10.1061/9780784485262.70" ext-link-type="DOI">10.1061/9780784485262.70</ext-link>, 2024.</mixed-citation></ref>
      <ref id="bib1.bib18"><label>18</label><mixed-citation>Darghiasi, P., Zamanian, M., Bhatta, S., and Shahandashti, M.: Enhanced road surface temperature prediction using random forest model and NWS weather forecast data, in: International Conference on Transportation and Development 2025, American Society of Civil Engineers, Reston, VA, USA, 286–298, <ext-link xlink:href="https://doi.org/10.1061/9780784486191.025" ext-link-type="DOI">10.1061/9780784486191.025</ext-link>, 2025.</mixed-citation></ref>
      <ref id="bib1.bib19"><label>19</label><mixed-citation>Diefenderfer, B. K., Al-Qadi, I. L., and Diefenderfer, S. D.: Model to predict pavement temperature profile: development and validation, J. Transp. Eng., 132, 162–167, <ext-link xlink:href="https://doi.org/10.1061/(ASCE)0733-947X(2006)132:2(162)" ext-link-type="DOI">10.1061/(ASCE)0733-947X(2006)132:2(162)</ext-link>, 2006.</mixed-citation></ref>
      <ref id="bib1.bib20"><label>20</label><mixed-citation>Divina, F., Gilson, A., Gómez-Vela, F., García Torres, M., and Torres, J. F.: Stacking ensemble learning for short-term electricity consumption forecasting, Energies, 11, 949, <ext-link xlink:href="https://doi.org/10.3390/en11040949" ext-link-type="DOI">10.3390/en11040949</ext-link>, 2018.</mixed-citation></ref>
      <ref id="bib1.bib21"><label>21</label><mixed-citation>Feng, T. and Feng, S.: A numerical model for predicting road surface temperature in the highway, Procedia Engineer., 37, 137–142, <ext-link xlink:href="https://doi.org/10.1016/j.proeng.2012.04.216" ext-link-type="DOI">10.1016/j.proeng.2012.04.216</ext-link>, 2012.</mixed-citation></ref>
      <ref id="bib1.bib22"><label>22</label><mixed-citation>Gedafa, D. S., Hossain, M., and Romanoschi, S. A.: Perpetual pavement temperature prediction model, Road Mater. Pavement Des., 15, 55–65, <ext-link xlink:href="https://doi.org/10.1080/14680629.2013.852610" ext-link-type="DOI">10.1080/14680629.2013.852610</ext-link>, 2014.</mixed-citation></ref>
      <ref id="bib1.bib23"><label>23</label><mixed-citation>Gelman, A. and Shalizi, C. R.: Philosophy and the practice of Bayesian statistics, Brit. J. Math. Stat. Psy., 66, 8–38, <ext-link xlink:href="https://doi.org/10.1111/j.2044-8317.2011.02037.x" ext-link-type="DOI">10.1111/j.2044-8317.2011.02037.x</ext-link>, 2013.</mixed-citation></ref>
      <ref id="bib1.bib24"><label>24</label><mixed-citation>Ghalandari, T., Shi, L., Sadeghi-Khanegah, F., Van den Bergh, W., and Vuye, C.: Utilizing artificial neural networks to predict the asphalt pavement profile temperature in western Europe, Case Stud. Constr. Mater., 18, e02130, <ext-link xlink:href="https://doi.org/10.1016/j.cscm.2023.e02130" ext-link-type="DOI">10.1016/j.cscm.2023.e02130</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib25"><label>25</label><mixed-citation> Glorot, X., Bordes, A., and Bengio, Y.: Deep sparse rectifier neural networks, in: Proceedings of the 14th International Conference on Artificial Intelligence and Statistics, Proceedings of Machine Learning Research (PMLR), 15, 315–323, 2011.</mixed-citation></ref>
      <ref id="bib1.bib26"><label>26</label><mixed-citation>Gui, J., Phelan, P. E., Kaloush, K. E., and Golden, J. S.:  Impact of pavement thermophysical properties on surface temperatures, J. Mater. Civil Eng., 19, 683–690, <ext-link xlink:href="https://doi.org/10.1061/(ASCE)0899-1561(2007)19:8(683)" ext-link-type="DOI">10.1061/(ASCE)0899-1561(2007)19:8(683)</ext-link>, 2007.</mixed-citation></ref>
      <ref id="bib1.bib27"><label>27</label><mixed-citation>Hassan, H. F., Al-Nuaimi, A. S., Taha, R., and Jafar, T. M.: Development of asphalt pavement temperature models for Oman, J. Eng. Res., 2, 32–42, <ext-link xlink:href="https://doi.org/10.24200/tjer.vol2iss1pp32-42" ext-link-type="DOI">10.24200/tjer.vol2iss1pp32-42</ext-link>, 2005.</mixed-citation></ref>
      <ref id="bib1.bib28"><label>28</label><mixed-citation>Hatamzad, M., Pinerez, G. C. P., and Casselgren, J.: Intelligent cost-effective winter road maintenance by predicting road surface temperature using machine learning techniques, Knowl.-Based Syst., 247, 108682, <ext-link xlink:href="https://doi.org/10.1016/j.knosys.2022.108682" ext-link-type="DOI">10.1016/j.knosys.2022.108682</ext-link>, 2022.</mixed-citation></ref>
      <ref id="bib1.bib29"><label>29</label><mixed-citation>He, K., Zhang, X., Ren, S., and Sun, J.: Deep residual learning for image recognition, in: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, IEEE, Piscataway, NJ, USA, 770–778, <ext-link xlink:href="https://doi.org/10.1109/CVPR.2016.90" ext-link-type="DOI">10.1109/CVPR.2016.90</ext-link>, 2016.</mixed-citation></ref>
      <ref id="bib1.bib30"><label>30</label><mixed-citation>Hermansson, Å.: Mathematical model for paved surface summer and winter temperature: comparison of calculated and measured temperatures, Cold Reg. Sci. Technol., 40, 1–17, <ext-link xlink:href="https://doi.org/10.1016/j.coldregions.2004.01.002" ext-link-type="DOI">10.1016/j.coldregions.2004.01.002</ext-link>, 2004.</mixed-citation></ref>
      <ref id="bib1.bib31"><label>31</label><mixed-citation>Hochreiter, S. and Schmidhuber, J.: Long short-term memory, Neural Comput., 9, 1735–1780, <ext-link xlink:href="https://doi.org/10.1162/neco.1997.9.8.1735" ext-link-type="DOI">10.1162/neco.1997.9.8.1735</ext-link>, 1997.</mixed-citation></ref>
      <ref id="bib1.bib32"><label>32</label><mixed-citation>Hoerl, A. E. and Kennard, R. W.: Ridge regression: Biased estimation for nonorthogonal problems, Technometrics, 12, 55–67, <ext-link xlink:href="https://doi.org/10.1080/00401706.1970.10488634" ext-link-type="DOI">10.1080/00401706.1970.10488634</ext-link>, 1970.</mixed-citation></ref>
      <ref id="bib1.bib33"><label>33</label><mixed-citation>Jing, C. and Zhang, J.: Prediction model for asphalt pavement temperature in high‐temperature season in Beijing, Adv. Civ. Eng., 2018, 1837952, <ext-link xlink:href="https://doi.org/10.1155/2018/1837952" ext-link-type="DOI">10.1155/2018/1837952</ext-link>, 2018.</mixed-citation></ref>
      <ref id="bib1.bib34"><label>34</label><mixed-citation>Joo, C., Park, H., Lim, J., Cho, H., and Kim, J.: Learning-based heat deflection temperature prediction and effect analysis in polypropylene composites using catboost and shapley additive explanations, Eng. Appl. Artif. Intel., 126, 106873, <ext-link xlink:href="https://doi.org/10.1016/j.engappai.2023.106873" ext-link-type="DOI">10.1016/j.engappai.2023.106873</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib35"><label>35</label><mixed-citation>Kangas, M., Heikinheimo, M., and Hippi, M.: RoadSurf: a behaviour system for predicting road weather and road surface conditions, Meteorol. Appl., 22, 544–553, <ext-link xlink:href="https://doi.org/10.1002/met.1486" ext-link-type="DOI">10.1002/met.1486</ext-link>, 2015.</mixed-citation></ref>
      <ref id="bib1.bib36"><label>36</label><mixed-citation>Karsisto, V., Nurmi, P., Kangas, M., Hippi, M., and Uppala, A.: Improving road weather model forecasts by adjusting the radiation input, Meteorol. Appl., 23, 503–513, <ext-link xlink:href="https://doi.org/10.1002/met.1574" ext-link-type="DOI">10.1002/met.1574</ext-link>, 2016.</mixed-citation></ref>
      <ref id="bib1.bib37"><label>37</label><mixed-citation>Kebede, Y. B., Yang, M. D., and Huang, C. W.: Real-time pavement temperature prediction through ensemble machine learning, Eng. Appl. Artif. Intel., 135, 108870, <ext-link xlink:href="https://doi.org/10.1016/j.engappai.2024.108870" ext-link-type="DOI">10.1016/j.engappai.2024.108870</ext-link>, 2024.</mixed-citation></ref>
      <ref id="bib1.bib38"><label>38</label><mixed-citation>Kingma, D. P. and Ba, J.: Adam: A method for stochastic optimization, arXiv [preprint], <ext-link xlink:href="https://doi.org/10.48550/arXiv.1412.6980" ext-link-type="DOI">10.48550/arXiv.1412.6980</ext-link>, 2014.</mixed-citation></ref>
      <ref id="bib1.bib39"><label>39</label><mixed-citation>Kršmanc, R., Slak, A. Š., and Demšar, J.: Statistical approach for forecasting road surface temperature, Meteorol. Appl., 20, 439–446, <ext-link xlink:href="https://doi.org/10.1002/met.1305" ext-link-type="DOI">10.1002/met.1305</ext-link>, 2013.</mixed-citation></ref>
      <ref id="bib1.bib40"><label>40</label><mixed-citation>Kuncheva, L. I. and Whitaker, C. J.: Measures of diversity in classifier ensembles and their relationship with the ensemble accuracy, Mach. Learn., 51, 181–207, <ext-link xlink:href="https://doi.org/10.1023/a:1022859003006" ext-link-type="DOI">10.1023/a:1022859003006</ext-link>, 2003.</mixed-citation></ref>
      <ref id="bib1.bib41"><label>41</label><mixed-citation>Li, W.: A hybrid method for winter road surface temperature prediction using improved LSTMs and stacking-based ensemble learning, Zenodo [code, data set], <ext-link xlink:href="https://doi.org/10.5281/zenodo.22020890" ext-link-type="DOI">10.5281/zenodo.22020890</ext-link>, 2026.</mixed-citation></ref>
      <ref id="bib1.bib42"><label>42</label><mixed-citation>Li, Y., Liu, L., and Sun, L.: Temperature predictions for asphalt pavement with thick asphalt layer, Constr. Build. Mater., 160, 802–809, <ext-link xlink:href="https://doi.org/10.1016/j.conbuildmat.2017.11.077" ext-link-type="DOI">10.1016/j.conbuildmat.2017.11.077</ext-link>, 2018.</mixed-citation></ref>
      <ref id="bib1.bib43"><label>43</label><mixed-citation>Li, Y., Chen, J., Dan, H., and Wang, H.:  Probability prediction of pavement surface low temperature in winter based on bayesian structural time series and neural network, Cold Reg. Sci. Technol., 194, 103434, <ext-link xlink:href="https://doi.org/10.1016/j.coldregions.2021.103434" ext-link-type="DOI">10.1016/j.coldregions.2021.103434</ext-link>, 2022.</mixed-citation></ref>
      <ref id="bib1.bib44"><label>44</label><mixed-citation>Lin, M., Chen, Q., and Yan, S.: Network in network, arXiv [preprint], <ext-link xlink:href="https://doi.org/10.48550/arXiv.1312.4400" ext-link-type="DOI">10.48550/arXiv.1312.4400</ext-link>, 2013.</mixed-citation></ref>
      <ref id="bib1.bib45"><label>45</label><mixed-citation>Liu, B., Yan, S., You, H., Dong, Y., Li, Y., Lang, J., and Gu, R.: Road surface temperature prediction based on gradient extreme learning machine boosting, Comput. Ind., 99, 294–302, <ext-link xlink:href="https://doi.org/10.1016/j.compind.2018.03.026" ext-link-type="DOI">10.1016/j.compind.2018.03.026</ext-link>, 2018.</mixed-citation></ref>
      <ref id="bib1.bib46"><label>46</label><mixed-citation>Lundberg, S. M. and Lee, S. I.: A unified approach to interpreting model predictions, Adv. Neur. In., arXiv [preprint], <ext-link xlink:href="https://doi.org/10.48550/arXiv.1705.07874" ext-link-type="DOI">10.48550/arXiv.1705.07874</ext-link>, 2017.</mixed-citation></ref>
      <ref id="bib1.bib47"><label>47</label><mixed-citation>Luo, X., Li, D., Yang, Y., and Zhang, S.: Spatiotemporal traffic flow prediction with KNN and LSTM, J. Adv. Transp., 2019, 4145353, <ext-link xlink:href="https://doi.org/10.1155/2019/4145353" ext-link-type="DOI">10.1155/2019/4145353</ext-link>, 2019.</mixed-citation></ref>
      <ref id="bib1.bib48"><label>48</label><mixed-citation>MacKay, D. J. C.: Bayesian interpolation, Neural Comput., 4, 415–447, <ext-link xlink:href="https://doi.org/10.1162/neco.1992.4.3.415" ext-link-type="DOI">10.1162/neco.1992.4.3.415</ext-link>, 1992.</mixed-citation></ref>
      <ref id="bib1.bib49"><label>49</label><mixed-citation>Maddu, R., Vanga, A. R., Sajja, J. K., Basha, G., and Shaik, R.: Prediction of land surface temperature of major coastal cities of India using bidirectional LSTM neural networks, J. Water Clim. Change, 12, 3801–3819, <ext-link xlink:href="https://doi.org/10.2166/wcc.2021.460" ext-link-type="DOI">10.2166/wcc.2021.460</ext-link>, 2021.</mixed-citation></ref>
      <ref id="bib1.bib50"><label>50</label><mixed-citation>Milad, A., Adwan, I., Majeed, S. A., Yusoff, N. I. M., Al-Ansari, N., and Yaseen, Z. M.: Emerging technologies of deep learning models development for pavement temperature prediction, IEEE Access, 9, 23840–23849, <ext-link xlink:href="https://doi.org/10.1109/ACCESS.2021.3056568" ext-link-type="DOI">10.1109/ACCESS.2021.3056568</ext-link>, 2021a.</mixed-citation></ref>
      <ref id="bib1.bib51"><label>51</label><mixed-citation>Milad, A. A., Adwan, I., Majeed, S. A., Memon, Z. A., Bilema, M., and Omar, H. A.: Development of a hybrid machine learning model for asphalt pavement temperature prediction, IEEE Access, 9, 158041–158056, <ext-link xlink:href="https://doi.org/10.1109/ACCESS.2021.3129979" ext-link-type="DOI">10.1109/ACCESS.2021.3129979</ext-link>, 2021b.</mixed-citation></ref>
      <ref id="bib1.bib52"><label>52</label><mixed-citation>Minhoto, M. J. C., Pais, J. C., Pereira, P. A., and Picado-Santos, L.: Predicting asphalt pavement temperature with a three-dimensional finite element method, Transp. Res. Rec., 1919, 96–110, <ext-link xlink:href="https://doi.org/10.1177/0361198105191900111" ext-link-type="DOI">10.1177/0361198105191900111</ext-link>, 2005.</mixed-citation></ref>
      <ref id="bib1.bib53"><label>53</label><mixed-citation>Molavi Nojumi, M., Huang, Y., Hashemian, L., and Bayat, A.: Application of machine learning for temperature prediction in a test road in Alberta, Int. J. Pavement Res. Technol., 15, 303–319, <ext-link xlink:href="https://doi.org/10.1007/s42947-021-00023-3" ext-link-type="DOI">10.1007/s42947-021-00023-3</ext-link>, 2022.</mixed-citation></ref>
      <ref id="bib1.bib54"><label>54</label><mixed-citation>Muñoz Sabater, J.: ERA5-Land hourly data from 1950 to present, Copernicus Climate Change Service (C3S) Climate Data Store (CDS) [data set], <ext-link xlink:href="https://doi.org/10.24381/cds.e2161bac" ext-link-type="DOI">10.24381/cds.e2161bac</ext-link>, 2019.</mixed-citation></ref>
      <ref id="bib1.bib55"><label>55</label><mixed-citation>Muñoz-Sabater, J., Dutra, E., Agustí-Panareda, A., Albergel, C., Arduini, G., Balsamo, G., Boussetta, S., Choulga, M., Harrigan, S., Hersbach, H., Martens, B., Miralles, D. G., Piles, M., Rodríguez-Fernández, N. J., Zsoter, E., Buontempo, C., and Thépaut, J.-N.: ERA5-Land: a state-of-the-art global reanalysis dataset for land applications, Earth Syst. Sci. Data, 13, 4349–4383, <ext-link xlink:href="https://doi.org/10.5194/essd-13-4349-2021" ext-link-type="DOI">10.5194/essd-13-4349-2021</ext-link>, 2021.</mixed-citation></ref>
      <ref id="bib1.bib56"><label>56</label><mixed-citation>Nantasai, B. and Nassiri, S.: Winter temperature prediction for near-surface depth of pervious concrete pavement, Int. J. Pavement Eng., 20, 820–829, <ext-link xlink:href="https://doi.org/10.1080/10298436.2017.1353389" ext-link-type="DOI">10.1080/10298436.2017.1353389</ext-link>, 2019.</mixed-citation></ref>
      <ref id="bib1.bib57"><label>57</label><mixed-citation>Nowrin, T. and Kwon, T. J.: Forecasting short-term road surface temperatures considering forecasting horizon and geographical attributes–an ANN-based approach, Cold Reg. Sci. Technol., 202, 103631, <ext-link xlink:href="https://doi.org/10.1016/j.coldregions.2022.103631" ext-link-type="DOI">10.1016/j.coldregions.2022.103631</ext-link>, 2022.</mixed-citation></ref>
      <ref id="bib1.bib58"><label>58</label><mixed-citation>Qin, Y. and Hiller, J. E.: Ways of formulating wind speed in heat convection significantly influencing pavement temperature prediction, Heat Mass Transfer, 49, 745–752, <ext-link xlink:href="https://doi.org/10.1007/s00231-013-1116-0" ext-link-type="DOI">10.1007/s00231-013-1116-0</ext-link>, 2013.</mixed-citation></ref>
      <ref id="bib1.bib59"><label>59</label><mixed-citation>Qin, Y., Zhang, X., Tan, K., and Wang, J.: A review on the influencing factors of pavement surface temperature, Environ. Sci. Pollut. R., 29, 67659–67674, <ext-link xlink:href="https://doi.org/10.1007/s11356-022-22295-3" ext-link-type="DOI">10.1007/s11356-022-22295-3</ext-link>, 2022.</mixed-citation></ref>
      <ref id="bib1.bib60"><label>60</label><mixed-citation>Qiu, X., Hong, H., Xu, W., Yang, Q., and Xiao, S.:  Surface temperature prediction of asphalt pavement based on APRIORI-GBDT, in: International Conference on Transportation and Development 2020, 200–212, <ext-link xlink:href="https://doi.org/10.1061/9780784483183.020" ext-link-type="DOI">10.1061/9780784483183.020</ext-link>, 2020.</mixed-citation></ref>
      <ref id="bib1.bib61"><label>61</label><mixed-citation>Reichstein, M., Camps-Valls, G., Stevens, B., Jung, M., Denzler, J., Carvalhais, N., and Prabhat, F.:  Deep learning and process understanding for data-driven Earth system science, Nature, 566, 195–204, <ext-link xlink:href="https://doi.org/10.1038/s41586-019-0912-1" ext-link-type="DOI">10.1038/s41586-019-0912-1</ext-link>, 2019.</mixed-citation></ref>
      <ref id="bib1.bib62"><label>62</label><mixed-citation>Rigabadi, A., Rezaei Zadeh Herozi, M., and Rezagholilou, A.: An attempt for development of pavements temperature prediction models based on remote sensing data and artificial neural network, Int. J. Pavement Eng., 23, 2912–2921, <ext-link xlink:href="https://doi.org/10.1080/10298436.2021.1873334" ext-link-type="DOI">10.1080/10298436.2021.1873334</ext-link>, 2022.</mixed-citation></ref>
      <ref id="bib1.bib63"><label>63</label><mixed-citation>Saliko, D., Ahmed, A., and Erlingsson, S.: Development and validation of a pavement temperature profile prediction model in a mechanistic-empirical design framework, Transp. Geotech., 40, 100976, <ext-link xlink:href="https://doi.org/10.1016/j.trgeo.2023.100976" ext-link-type="DOI">10.1016/j.trgeo.2023.100976</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib64"><label>64</label><mixed-citation> Schindler, A. K., Ruiz, J. M., Rasmussen, R. O., Chang, G. K., and Wathne, L. G.: Concrete pavement temperature prediction and case studies with the FHWA HIPERPAV models, Cem. Concr. Compos., 26, 463–471, https://doi.org/10.1016/s0958-9465(03)00075-1, 2004.</mixed-citation></ref>
      <ref id="bib1.bib65"><label>65</label><mixed-citation>Schuster, M. and Paliwal, K. K.: Bidirectional recurrent neural networks, IEEE T. Signal Proces., 45, 2673–2681, <ext-link xlink:href="https://doi.org/10.1109/78.650093" ext-link-type="DOI">10.1109/78.650093</ext-link>, 1997.</mixed-citation></ref>
      <ref id="bib1.bib66"><label>66</label><mixed-citation>Shao, J. and Lister, P. J.: An automated nowcasting model of road surface temperature and state for winter road maintenance, J. Appl. Meteorol., 35, 1352–1361, <ext-link xlink:href="https://doi.org/10.1175/1520-0450(1996)035&lt;1352:AANMOR&gt;2.0.CO;2" ext-link-type="DOI">10.1175/1520-0450(1996)035&lt;1352:AANMOR&gt;2.0.CO;2</ext-link>, 1996.</mixed-citation></ref>
      <ref id="bib1.bib67"><label>67</label><mixed-citation>Song, P., Che, J., and Guo, T.: Climatic characteristics and SVM forecast model of subfreezing road temperature on expressways, J. Mar. Meteorol., 29, 56–64, <ext-link xlink:href="https://doi.org/10.19513/j.cnki.issn2096-3599.2023.03.008" ext-link-type="DOI">10.19513/j.cnki.issn2096-3599.2023.03.008</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib68"><label>68</label><mixed-citation>Spearman, C.: The proof and measurement of association between two things, in: Studies in Individual Differences: The Search for Intelligence, edited by: Jenkins, J. J. and Paterson, D. G., Appleton-Century-Crofts, New York, NY, USA, 45–58, <ext-link xlink:href="https://doi.org/10.1037/11491-005" ext-link-type="DOI">10.1037/11491-005</ext-link>, 1961.</mixed-citation></ref>
      <ref id="bib1.bib69"><label>69</label><mixed-citation>Tabrizi, S. E., Xiao, K., Thé, J. V. G., Saad, M., Farghaly, H., Yang, S. X., and Gharabaghi, B.: Hourly road pavement surface temperature forecasting using deep learning models, J. Hydrol., 603, 126877, <ext-link xlink:href="https://doi.org/10.1016/j.jhydrol.2021.126877" ext-link-type="DOI">10.1016/j.jhydrol.2021.126877</ext-link>, 2021.</mixed-citation></ref>
      <ref id="bib1.bib70"><label>70</label><mixed-citation>Tao, R., Peng, R., Wang, H., Wang, J., and Qiao, J.:  Temperature and humidity prediction of mountain highway tunnel entrance road surface based on improved Bi-LSTM neural network, Evol. Syst., 15, 691–702, <ext-link xlink:href="https://doi.org/10.1007/s12530-023-09538-5" ext-link-type="DOI">10.1007/s12530-023-09538-5</ext-link>, 2024.</mixed-citation></ref>
      <ref id="bib1.bib71"><label>71</label><mixed-citation>Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A. N., Kaiser, Ł., and Polosukhin, I.:  Attention is all you need, Adv. Neur. In., 30, <ext-link xlink:href="https://doi.org/10.1201/9781003561460-19" ext-link-type="DOI">10.1201/9781003561460-19</ext-link>, 2017.</mixed-citation></ref>
      <ref id="bib1.bib72"><label>72</label><mixed-citation>Wang, D.: Simplified analytical approach to predicting asphalt pavement temperature, J. Mater. Civil Eng., 27, 04015043, <ext-link xlink:href="https://doi.org/10.1061/(ASCE)MT.1943-5533.0001301" ext-link-type="DOI">10.1061/(ASCE)MT.1943-5533.0001301</ext-link>, 2015.</mixed-citation></ref>
      <ref id="bib1.bib73"><label>73</label><mixed-citation>Wang, D., Roesler, J. R., and Guo, D. Z.: Analytical approach to predicting temperature fields in multilayered pavement systems, J. Eng. Mech., 135, 334–344, <ext-link xlink:href="https://doi.org/10.1061/(ASCE)0733-9399(2009)135:4(334)" ext-link-type="DOI">10.1061/(ASCE)0733-9399(2009)135:4(334)</ext-link>, 2009.</mixed-citation></ref>
      <ref id="bib1.bib74"><label>74</label><mixed-citation>Wang, X., Pan, P., and Li, J.: Real-time measurement on dynamic temperature variation of asphalt pavement using machine learning, Measurement, 207, 112413, <ext-link xlink:href="https://doi.org/10.1016/j.measurement.2022.112413" ext-link-type="DOI">10.1016/j.measurement.2022.112413</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib75"><label>75</label><mixed-citation>Wolpert, D. H.: Stacked generalization, Neural Netw., 5, 241–259, <ext-link xlink:href="https://doi.org/10.1016/S0893-6080(05)80023-1" ext-link-type="DOI">10.1016/S0893-6080(05)80023-1</ext-link>, 1992. </mixed-citation></ref>
      <ref id="bib1.bib76"><label>76</label><mixed-citation>Yang, C. H., Yun, D. G., Kim, J. G., Lee, G., and Kim, S. B.: Machine learning approaches to estimate road surface temperature variation along road section in real-time for winter operation, Int. J. Intell. Transp. Syst. Res., 18, 343–355, <ext-link xlink:href="https://doi.org/10.1007/s13177-019-00198-x" ext-link-type="DOI">10.1007/s13177-019-00198-x</ext-link>, 2020.</mixed-citation></ref>
      <ref id="bib1.bib77"><label>77</label><mixed-citation>Yin, Z., Hadzimustafic, J., Kann, A., and Wang, Y.:  On statistical nowcasting of road surface temperature, Meteorol. Appl., 26, 1–13, <ext-link xlink:href="https://doi.org/10.1002/met.1737" ext-link-type="DOI">10.1002/met.1737</ext-link>, 2019.</mixed-citation></ref>
      <ref id="bib1.bib78"><label>78</label><mixed-citation>Yuan, J., Cheng, H., Sun, L., Cao, Y., Yang, R., Jin, T., and Li, M.:  Cross-Regional Pavement Temperature Prediction Using Transfer Learning and Random Forest, Appl. Sci., 15, 7436, <ext-link xlink:href="https://doi.org/10.3390/app15137436" ext-link-type="DOI">10.3390/app15137436</ext-link>, 2025.</mixed-citation></ref>
      <ref id="bib1.bib79"><label>79</label><mixed-citation>Zhang, M., Guo, H., Li, J. Y., Li, L., and Zhu, F.:  A deep learning approach for enhanced real-time prediction of winter road surface temperatures in high-altitude mountain areas, Promet, 36, 958–972, <ext-link xlink:href="https://doi.org/10.7307/ptt.v36i5.541" ext-link-type="DOI">10.7307/ptt.v36i5.541</ext-link>, 2024.</mixed-citation></ref>
      <ref id="bib1.bib80"><label>80</label><mixed-citation>Zhang, N., Mao, T., Chen, H., Lv, L., Wang, Y., and Yan, Y.:  Temperature prediction for expressway pavement icing in winter based on XGBoost–LSTNet variable weight combination model, J. Transp. Eng. A-Syst., 149, 04023062, <ext-link xlink:href="https://doi.org/10.1061/JTEPBS.TEENG-7918" ext-link-type="DOI">10.1061/JTEPBS.TEENG-7918</ext-link>, 2023.</mixed-citation></ref>
      <ref id="bib1.bib81"><label>81</label><mixed-citation>Zhao, X., Shen, A., and Ma, B.: Temperature response of asphalt pavement to low temperatures and large temperature differences, Int. J. Pavement Eng., 21, 49–62, <ext-link xlink:href="https://doi.org/10.1080/10298436.2018.1435877" ext-link-type="DOI">10.1080/10298436.2018.1435877</ext-link>, 2020.</mixed-citation></ref>

  </ref-list></back>
    <!--<article-title-html>A hybrid method for winter road surface temperature prediction using improved LSTMs and stacking-based ensemble learning</article-title-html>
<abstract-html/>
<ref-html id="bib1.bib1"><label>1</label><mixed-citation>
      
Abo-Hashema, M. A.: Modeling pavement temperature prediction using artificial neural networks, in: Airfield and highway pavement 2013: Sustainable and efficient pavements, American Society of Civil Engineers, Reston, VA, USA, 490–505,  <a href="https://doi.org/10.1061/9780784413005.039" target="_blank">https://doi.org/10.1061/9780784413005.039</a>, 2013.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib2"><label>2</label><mixed-citation>
      
Adwan, I., Milad, A., Memon, Z. A., Widyatmoko, I., Zanuri, N. A., Memon, N. A., and Yusoff, N. I. M.:  Asphalt pavement temperature prediction models: A review, Appl. Sci., 11, 3794, <a href="https://doi.org/10.3390/app11093794" target="_blank">https://doi.org/10.3390/app11093794</a>, 2021.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib3"><label>3</label><mixed-citation>
      
Asefzadeh, A., Hashemian, L., and Bayat, A.: Development of statistical temperature prediction models for a test road in Edmonton, Alberta, Canada, Int. J. Pavement Res. Technol., 10, 369–382, <a href="https://doi.org/10.1016/j.ijprt.2017.05.003" target="_blank">https://doi.org/10.1016/j.ijprt.2017.05.003</a>, 2017.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib4"><label>4</label><mixed-citation>
      
Athukorallage, B., Senadheera, S., and James, D.: Temporal and spatial temperature predictions for flexible pavement layers using numerical thermal analysis and verified with large datasets, Case Stud. Constr. Mater., 18, e02008, <a href="https://doi.org/10.1016/j.cscm.2023.e02008" target="_blank">https://doi.org/10.1016/j.cscm.2023.e02008</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib5"><label>5</label><mixed-citation>
      
Ayasrah, U. B., Tashman, L., AlOmari, A., and Asi, I.: Development of a temperature prediction model for flexible pavement structures, Case Stud. Constr. Mater., 18, e01697, <a href="https://doi.org/10.1016/j.cscm.2022.e01697" target="_blank">https://doi.org/10.1016/j.cscm.2022.e01697</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib6"><label>6</label><mixed-citation>
      
Ba, J. L., Kiros, J. R., and Hinton, G. E.: Layer normalization, arXiv [preprint], <a href="https://doi.org/10.48550/arXiv.1607.06450" target="_blank">https://doi.org/10.48550/arXiv.1607.06450</a>, 2016.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib7"><label>7</label><mixed-citation>
      
Bai, S., Yang, W., Zhang, M., Liu, D., Li, W., and Zhou, L.: Attention-based BiLSTM model for pavement temperature prediction of asphalt pavement in winter, Atmosphere, 13, 1524, <a href="https://doi.org/10.3390/atmos13091524" target="_blank">https://doi.org/10.3390/atmos13091524</a>, 2022.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib8"><label>8</label><mixed-citation>
      
Ben Taieb, S., Bontempi, G., Atiya, A. F., and Sorjamaa, A.: A review and comparison of strategies for multi-step ahead time series forecasting based on the NN5 forecasting competition, Expert Syst. Appl., 39, 7067–7083, <a href="https://doi.org/10.1016/j.eswa.2012.01.039" target="_blank">https://doi.org/10.1016/j.eswa.2012.01.039</a>, 2012.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib9"><label>9</label><mixed-citation>
      
Bishop, C. M., Nasrabadi, N. M.: Pattern recognition and machine learning, New York, Springer, <a href="https://doi.org/10.1080/15228053.2019.1632410" target="_blank">https://doi.org/10.1080/15228053.2019.1632410</a>, 2006.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib10"><label>10</label><mixed-citation>
      
Breiman, L.: Bagging predictors, Mach. Learn., 24, 123–140, <a href="https://doi.org/10.1007/BF00058655" target="_blank">https://doi.org/10.1007/BF00058655</a>, 1996.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib11"><label>11</label><mixed-citation>
      
Chen, J., Wang, H., and Xie, P.: Pavement temperature prediction: Theoretical models and critical affecting factors, Appl. Therm. Eng., 158, 113755, <a href="https://doi.org/10.1016/j.applthermaleng.2019.113755" target="_blank">https://doi.org/10.1016/j.applthermaleng.2019.113755</a>, 2019.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib12"><label>12</label><mixed-citation>
      
Cheng, H., Liu, J., Sun, L., and Liu, L.: Critical position of fatigue damage within asphalt pavement considering temperature and strain distribution, Int. J. Pavement Eng., 22, 1773–1784, <a href="https://doi.org/10.1080/10298436.2020.1724288" target="_blank">https://doi.org/10.1080/10298436.2020.1724288</a>, 2021.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib13"><label>13</label><mixed-citation>
      
China Meteorological Administration: Grade of highway traffic high-impact weather warning, QX/T 414-2018, China Meteorological Press, Beijing, China, 2018.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib14"><label>14</label><mixed-citation>
      
Crevier, L.-P. and Delage, Y.: METRo: A new model for road-condition forecasting in Canada, J. Appl. Meteorol., 40, 2026–2037, <a href="https://doi.org/10.1175/1520-0450(2001)040&lt;2026:MANMFR&gt;2.0.CO;2" target="_blank">https://doi.org/10.1175/1520-0450(2001)040&lt;2026:MANMFR&gt;2.0.CO;2</a>, 2001.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib15"><label>15</label><mixed-citation>
      
Dai, B., Yang, W., Ji, X., and Zhou, L.: An ensemble deep learning model for short-term road surface temperature prediction, J. Transp. Eng. B-Pavements, 149, 04022067, <a href="https://doi.org/10.1061/JPEODX.PVENG-1192" target="_blank">https://doi.org/10.1061/JPEODX.PVENG-1192</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib16"><label>16</label><mixed-citation>
      
Darghiasi, P., Baral, A., Mattingly, S., and Shahandashti, M.: Estimation of road surface temperature using NOAA gridded forecast weather data for snowplow operations management, J. Cold Reg. Eng., 37, 04023018, <a href="https://doi.org/10.1061/JCRGEI.CRENG-691" target="_blank">https://doi.org/10.1061/JCRGEI.CRENG-691</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib17"><label>17</label><mixed-citation>
      
Darghiasi, P., Zamanian, M., and Shahandashti, M.: Enhancing Winter Maintenance Decision Making through Deep Learning-Based Road Surface Temperature Estimation, in: Construction Res. Congr. 2024, ASCE, 701–711, <a href="https://doi.org/10.1061/9780784485262.70" target="_blank">https://doi.org/10.1061/9780784485262.70</a>, 2024.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib18"><label>18</label><mixed-citation>
      
Darghiasi, P., Zamanian, M., Bhatta, S., and Shahandashti, M.: Enhanced road surface temperature prediction using random forest model and NWS weather forecast data, in: International Conference on Transportation and Development 2025, American Society of Civil Engineers, Reston, VA, USA, 286–298, <a href="https://doi.org/10.1061/9780784486191.025" target="_blank">https://doi.org/10.1061/9780784486191.025</a>, 2025.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib19"><label>19</label><mixed-citation>
      
Diefenderfer, B. K., Al-Qadi, I. L., and Diefenderfer, S. D.: Model to predict pavement temperature profile: development and validation, J. Transp. Eng., 132, 162–167, <a href="https://doi.org/10.1061/(ASCE)0733-947X(2006)132:2(162)" target="_blank">https://doi.org/10.1061/(ASCE)0733-947X(2006)132:2(162)</a>, 2006.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib20"><label>20</label><mixed-citation>
      
Divina, F., Gilson, A., Gómez-Vela, F., García Torres, M., and Torres, J. F.: Stacking ensemble learning for short-term electricity consumption forecasting, Energies, 11, 949, <a href="https://doi.org/10.3390/en11040949" target="_blank">https://doi.org/10.3390/en11040949</a>, 2018.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib21"><label>21</label><mixed-citation>
      
Feng, T. and Feng, S.: A numerical model for predicting road surface temperature in the highway, Procedia Engineer., 37, 137–142, <a href="https://doi.org/10.1016/j.proeng.2012.04.216" target="_blank">https://doi.org/10.1016/j.proeng.2012.04.216</a>, 2012.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib22"><label>22</label><mixed-citation>
      
Gedafa, D. S., Hossain, M., and Romanoschi, S. A.: Perpetual pavement temperature prediction model, Road Mater. Pavement Des., 15, 55–65, <a href="https://doi.org/10.1080/14680629.2013.852610" target="_blank">https://doi.org/10.1080/14680629.2013.852610</a>, 2014.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib23"><label>23</label><mixed-citation>
      
Gelman, A. and Shalizi, C. R.: Philosophy and the practice of Bayesian statistics, Brit. J. Math. Stat. Psy., 66, 8–38, <a href="https://doi.org/10.1111/j.2044-8317.2011.02037.x" target="_blank">https://doi.org/10.1111/j.2044-8317.2011.02037.x</a>, 2013.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib24"><label>24</label><mixed-citation>
      
Ghalandari, T., Shi, L., Sadeghi-Khanegah, F., Van den Bergh, W., and Vuye, C.: Utilizing artificial neural networks to predict the asphalt pavement profile temperature in western Europe, Case Stud. Constr. Mater., 18, e02130, <a href="https://doi.org/10.1016/j.cscm.2023.e02130" target="_blank">https://doi.org/10.1016/j.cscm.2023.e02130</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib25"><label>25</label><mixed-citation>
      
Glorot, X., Bordes, A., and Bengio, Y.: Deep sparse rectifier neural networks, in: Proceedings of the 14th International Conference on Artificial Intelligence and Statistics, Proceedings of Machine Learning Research (PMLR), 15, 315–323, 2011.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib26"><label>26</label><mixed-citation>
      
Gui, J., Phelan, P. E., Kaloush, K. E., and Golden, J. S.:  Impact of pavement thermophysical properties on surface temperatures, J. Mater. Civil Eng., 19, 683–690, <a href="https://doi.org/10.1061/(ASCE)0899-1561(2007)19:8(683)" target="_blank">https://doi.org/10.1061/(ASCE)0899-1561(2007)19:8(683)</a>, 2007.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib27"><label>27</label><mixed-citation>
      
Hassan, H. F., Al-Nuaimi, A. S., Taha, R., and Jafar, T. M.: Development of asphalt pavement temperature models for Oman, J. Eng. Res., 2, 32–42, <a href="https://doi.org/10.24200/tjer.vol2iss1pp32-42" target="_blank">https://doi.org/10.24200/tjer.vol2iss1pp32-42</a>, 2005.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib28"><label>28</label><mixed-citation>
      
Hatamzad, M., Pinerez, G. C. P., and Casselgren, J.: Intelligent cost-effective winter road maintenance by predicting road surface temperature using machine learning techniques, Knowl.-Based Syst., 247, 108682, <a href="https://doi.org/10.1016/j.knosys.2022.108682" target="_blank">https://doi.org/10.1016/j.knosys.2022.108682</a>, 2022.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib29"><label>29</label><mixed-citation>
      
He, K., Zhang, X., Ren, S., and Sun, J.: Deep residual learning for image recognition, in: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition, IEEE, Piscataway, NJ, USA, 770–778, <a href="https://doi.org/10.1109/CVPR.2016.90" target="_blank">https://doi.org/10.1109/CVPR.2016.90</a>, 2016.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib30"><label>30</label><mixed-citation>
      
Hermansson, Å.: Mathematical model for paved surface summer and winter temperature: comparison of calculated and measured temperatures, Cold Reg. Sci. Technol., 40, 1–17, <a href="https://doi.org/10.1016/j.coldregions.2004.01.002" target="_blank">https://doi.org/10.1016/j.coldregions.2004.01.002</a>, 2004.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib31"><label>31</label><mixed-citation>
      
Hochreiter, S. and Schmidhuber, J.: Long short-term memory, Neural Comput., 9, 1735–1780, <a href="https://doi.org/10.1162/neco.1997.9.8.1735" target="_blank">https://doi.org/10.1162/neco.1997.9.8.1735</a>, 1997.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib32"><label>32</label><mixed-citation>
      
Hoerl, A. E. and Kennard, R. W.: Ridge regression: Biased estimation for nonorthogonal problems, Technometrics, 12, 55–67, <a href="https://doi.org/10.1080/00401706.1970.10488634" target="_blank">https://doi.org/10.1080/00401706.1970.10488634</a>, 1970.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib33"><label>33</label><mixed-citation>
      
Jing, C. and Zhang, J.: Prediction model for asphalt pavement temperature in high‐temperature season in Beijing, Adv. Civ. Eng., 2018, 1837952, <a href="https://doi.org/10.1155/2018/1837952" target="_blank">https://doi.org/10.1155/2018/1837952</a>, 2018.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib34"><label>34</label><mixed-citation>
      
Joo, C., Park, H., Lim, J., Cho, H., and Kim, J.: Learning-based heat deflection temperature prediction and effect analysis in polypropylene composites using catboost and shapley additive explanations, Eng. Appl. Artif. Intel., 126, 106873, <a href="https://doi.org/10.1016/j.engappai.2023.106873" target="_blank">https://doi.org/10.1016/j.engappai.2023.106873</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib35"><label>35</label><mixed-citation>
      
Kangas, M., Heikinheimo, M., and Hippi, M.: RoadSurf: a behaviour system for predicting road weather and road surface conditions, Meteorol. Appl., 22, 544–553, <a href="https://doi.org/10.1002/met.1486" target="_blank">https://doi.org/10.1002/met.1486</a>, 2015.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib36"><label>36</label><mixed-citation>
      
Karsisto, V., Nurmi, P., Kangas, M., Hippi, M., and Uppala, A.: Improving road weather model forecasts by adjusting the radiation input, Meteorol. Appl., 23, 503–513, <a href="https://doi.org/10.1002/met.1574" target="_blank">https://doi.org/10.1002/met.1574</a>, 2016.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib37"><label>37</label><mixed-citation>
      
Kebede, Y. B., Yang, M. D., and Huang, C. W.: Real-time pavement temperature prediction through ensemble machine learning, Eng. Appl. Artif. Intel., 135, 108870, <a href="https://doi.org/10.1016/j.engappai.2024.108870" target="_blank">https://doi.org/10.1016/j.engappai.2024.108870</a>, 2024.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib38"><label>38</label><mixed-citation>
      
Kingma, D. P. and Ba, J.: Adam: A method for stochastic optimization, arXiv [preprint], <a href="https://doi.org/10.48550/arXiv.1412.6980" target="_blank">https://doi.org/10.48550/arXiv.1412.6980</a>, 2014.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib39"><label>39</label><mixed-citation>
      
Kršmanc, R., Slak, A. Š., and Demšar, J.: Statistical approach for forecasting road surface temperature, Meteorol. Appl., 20, 439–446, <a href="https://doi.org/10.1002/met.1305" target="_blank">https://doi.org/10.1002/met.1305</a>, 2013.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib40"><label>40</label><mixed-citation>
      
Kuncheva, L. I. and Whitaker, C. J.: Measures of diversity in classifier ensembles and their relationship with the ensemble accuracy, Mach. Learn., 51, 181–207, <a href="https://doi.org/10.1023/a:1022859003006" target="_blank">https://doi.org/10.1023/a:1022859003006</a>, 2003.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib41"><label>41</label><mixed-citation>
      
Li, W.: A hybrid method for winter road surface temperature prediction using improved LSTMs and stacking-based ensemble learning, Zenodo [code, data set], <a href="https://doi.org/10.5281/zenodo.22020890" target="_blank">https://doi.org/10.5281/zenodo.22020890</a>, 2026.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib42"><label>42</label><mixed-citation>
      
Li, Y., Liu, L., and Sun, L.: Temperature predictions for asphalt pavement with thick asphalt layer, Constr. Build. Mater., 160, 802–809, <a href="https://doi.org/10.1016/j.conbuildmat.2017.11.077" target="_blank">https://doi.org/10.1016/j.conbuildmat.2017.11.077</a>, 2018.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib43"><label>43</label><mixed-citation>
      
Li, Y., Chen, J., Dan, H., and Wang, H.:  Probability prediction of pavement surface low temperature in winter based on bayesian structural time series and neural network, Cold Reg. Sci. Technol., 194, 103434, <a href="https://doi.org/10.1016/j.coldregions.2021.103434" target="_blank">https://doi.org/10.1016/j.coldregions.2021.103434</a>, 2022.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib44"><label>44</label><mixed-citation>
      
Lin, M., Chen, Q., and Yan, S.: Network in network, arXiv [preprint], <a href="https://doi.org/10.48550/arXiv.1312.4400" target="_blank">https://doi.org/10.48550/arXiv.1312.4400</a>, 2013.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib45"><label>45</label><mixed-citation>
      
Liu, B., Yan, S., You, H., Dong, Y., Li, Y., Lang, J., and Gu, R.: Road surface temperature prediction based on gradient extreme learning machine boosting, Comput. Ind., 99, 294–302, <a href="https://doi.org/10.1016/j.compind.2018.03.026" target="_blank">https://doi.org/10.1016/j.compind.2018.03.026</a>, 2018.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib46"><label>46</label><mixed-citation>
      
Lundberg, S. M. and Lee, S. I.: A unified approach to interpreting model predictions, Adv. Neur. In., arXiv [preprint], <a href="https://doi.org/10.48550/arXiv.1705.07874" target="_blank">https://doi.org/10.48550/arXiv.1705.07874</a>, 2017.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib47"><label>47</label><mixed-citation>
      
Luo, X., Li, D., Yang, Y., and Zhang, S.: Spatiotemporal traffic flow prediction with KNN and LSTM, J. Adv. Transp., 2019, 4145353, <a href="https://doi.org/10.1155/2019/4145353" target="_blank">https://doi.org/10.1155/2019/4145353</a>, 2019.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib48"><label>48</label><mixed-citation>
      
MacKay, D. J. C.: Bayesian interpolation, Neural Comput., 4, 415–447, <a href="https://doi.org/10.1162/neco.1992.4.3.415" target="_blank">https://doi.org/10.1162/neco.1992.4.3.415</a>, 1992.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib49"><label>49</label><mixed-citation>
      
Maddu, R., Vanga, A. R., Sajja, J. K., Basha, G., and Shaik, R.: Prediction of land surface temperature of major coastal cities of India using bidirectional LSTM neural networks, J. Water Clim. Change, 12, 3801–3819, <a href="https://doi.org/10.2166/wcc.2021.460" target="_blank">https://doi.org/10.2166/wcc.2021.460</a>, 2021.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib50"><label>50</label><mixed-citation>
      
Milad, A., Adwan, I., Majeed, S. A., Yusoff, N. I. M., Al-Ansari, N., and Yaseen, Z. M.: Emerging technologies of deep learning models development for pavement temperature prediction, IEEE Access, 9, 23840–23849, <a href="https://doi.org/10.1109/ACCESS.2021.3056568" target="_blank">https://doi.org/10.1109/ACCESS.2021.3056568</a>, 2021a.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib51"><label>51</label><mixed-citation>
      
Milad, A. A., Adwan, I., Majeed, S. A., Memon, Z. A., Bilema, M., and Omar, H. A.: Development of a hybrid machine learning model for asphalt pavement temperature prediction, IEEE Access, 9, 158041–158056, <a href="https://doi.org/10.1109/ACCESS.2021.3129979" target="_blank">https://doi.org/10.1109/ACCESS.2021.3129979</a>, 2021b.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib52"><label>52</label><mixed-citation>
      
Minhoto, M. J. C., Pais, J. C., Pereira, P. A., and Picado-Santos, L.: Predicting asphalt pavement temperature with a three-dimensional finite element method, Transp. Res. Rec., 1919, 96–110, <a href="https://doi.org/10.1177/0361198105191900111" target="_blank">https://doi.org/10.1177/0361198105191900111</a>, 2005.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib53"><label>53</label><mixed-citation>
      
Molavi Nojumi, M., Huang, Y., Hashemian, L., and Bayat, A.: Application of machine learning for temperature prediction in a test road in Alberta, Int. J. Pavement Res. Technol., 15, 303–319, <a href="https://doi.org/10.1007/s42947-021-00023-3" target="_blank">https://doi.org/10.1007/s42947-021-00023-3</a>, 2022.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib54"><label>54</label><mixed-citation>
      
Muñoz Sabater, J.: ERA5-Land hourly data from 1950 to present, Copernicus Climate Change Service (C3S) Climate Data Store (CDS) [data set], <a href="https://doi.org/10.24381/cds.e2161bac" target="_blank">https://doi.org/10.24381/cds.e2161bac</a>, 2019.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib55"><label>55</label><mixed-citation>
      
Muñoz-Sabater, J., Dutra, E., Agustí-Panareda, A., Albergel, C., Arduini, G., Balsamo, G., Boussetta, S., Choulga, M., Harrigan, S., Hersbach, H., Martens, B., Miralles, D. G., Piles, M., Rodríguez-Fernández, N. J., Zsoter, E., Buontempo, C., and Thépaut, J.-N.: ERA5-Land: a state-of-the-art global reanalysis dataset for land applications, Earth Syst. Sci. Data, 13, 4349–4383, <a href="https://doi.org/10.5194/essd-13-4349-2021" target="_blank">https://doi.org/10.5194/essd-13-4349-2021</a>, 2021.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib56"><label>56</label><mixed-citation>
      
Nantasai, B. and Nassiri, S.: Winter temperature prediction for near-surface depth of pervious concrete pavement, Int. J. Pavement Eng., 20, 820–829, <a href="https://doi.org/10.1080/10298436.2017.1353389" target="_blank">https://doi.org/10.1080/10298436.2017.1353389</a>, 2019.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib57"><label>57</label><mixed-citation>
      
Nowrin, T. and Kwon, T. J.: Forecasting short-term road surface temperatures considering forecasting horizon and geographical attributes–an ANN-based approach, Cold Reg. Sci. Technol., 202, 103631, <a href="https://doi.org/10.1016/j.coldregions.2022.103631" target="_blank">https://doi.org/10.1016/j.coldregions.2022.103631</a>, 2022.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib58"><label>58</label><mixed-citation>
      
Qin, Y. and Hiller, J. E.: Ways of formulating wind speed in heat convection significantly influencing pavement temperature prediction, Heat Mass Transfer, 49, 745–752, <a href="https://doi.org/10.1007/s00231-013-1116-0" target="_blank">https://doi.org/10.1007/s00231-013-1116-0</a>, 2013.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib59"><label>59</label><mixed-citation>
      
Qin, Y., Zhang, X., Tan, K., and Wang, J.: A review on the influencing factors of pavement surface temperature, Environ. Sci. Pollut. R., 29, 67659–67674, <a href="https://doi.org/10.1007/s11356-022-22295-3" target="_blank">https://doi.org/10.1007/s11356-022-22295-3</a>, 2022.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib60"><label>60</label><mixed-citation>
      
Qiu, X., Hong, H., Xu, W., Yang, Q., and Xiao, S.:  Surface temperature prediction of asphalt pavement based on APRIORI-GBDT, in: International Conference on Transportation and Development 2020, 200–212, <a href="https://doi.org/10.1061/9780784483183.020" target="_blank">https://doi.org/10.1061/9780784483183.020</a>, 2020.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib61"><label>61</label><mixed-citation>
      
Reichstein, M., Camps-Valls, G., Stevens, B., Jung, M., Denzler, J., Carvalhais, N., and Prabhat, F.:  Deep learning and process understanding for data-driven Earth system science, Nature, 566, 195–204, <a href="https://doi.org/10.1038/s41586-019-0912-1" target="_blank">https://doi.org/10.1038/s41586-019-0912-1</a>, 2019.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib62"><label>62</label><mixed-citation>
      
Rigabadi, A., Rezaei Zadeh Herozi, M., and Rezagholilou, A.: An attempt for development of pavements temperature prediction models based on remote sensing data and artificial neural network, Int. J. Pavement Eng., 23, 2912–2921, <a href="https://doi.org/10.1080/10298436.2021.1873334" target="_blank">https://doi.org/10.1080/10298436.2021.1873334</a>, 2022.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib63"><label>63</label><mixed-citation>
      
Saliko, D., Ahmed, A., and Erlingsson, S.: Development and validation of a pavement temperature profile prediction model in a mechanistic-empirical design framework, Transp. Geotech., 40, 100976, <a href="https://doi.org/10.1016/j.trgeo.2023.100976" target="_blank">https://doi.org/10.1016/j.trgeo.2023.100976</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib64"><label>64</label><mixed-citation>
      
Schindler, A. K., Ruiz, J. M., Rasmussen, R. O., Chang, G. K., and Wathne, L. G.: Concrete pavement temperature prediction and case studies with the FHWA HIPERPAV models, Cem. Concr. Compos., 26, 463–471, https://doi.org/10.1016/s0958-9465(03)00075-1, 2004.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib65"><label>65</label><mixed-citation>
      
Schuster, M. and Paliwal, K. K.: Bidirectional recurrent neural networks, IEEE T. Signal Proces., 45, 2673–2681, <a href="https://doi.org/10.1109/78.650093" target="_blank">https://doi.org/10.1109/78.650093</a>, 1997.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib66"><label>66</label><mixed-citation>
      
Shao, J. and Lister, P. J.: An automated nowcasting model of road surface temperature and state for winter road maintenance, J. Appl. Meteorol., 35, 1352–1361, <a href="https://doi.org/10.1175/1520-0450(1996)035&lt;1352:AANMOR&gt;2.0.CO;2" target="_blank">https://doi.org/10.1175/1520-0450(1996)035&lt;1352:AANMOR&gt;2.0.CO;2</a>, 1996.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib67"><label>67</label><mixed-citation>
      
Song, P., Che, J., and Guo, T.: Climatic characteristics and SVM forecast model of subfreezing road temperature on expressways, J. Mar. Meteorol., 29, 56–64, <a href="https://doi.org/10.19513/j.cnki.issn2096-3599.2023.03.008" target="_blank">https://doi.org/10.19513/j.cnki.issn2096-3599.2023.03.008</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib68"><label>68</label><mixed-citation>
      
Spearman, C.: The proof and measurement of association between two things, in: Studies in Individual Differences: The Search for Intelligence, edited by: Jenkins, J. J. and Paterson, D. G., Appleton-Century-Crofts, New York, NY, USA, 45–58, <a href="https://doi.org/10.1037/11491-005" target="_blank">https://doi.org/10.1037/11491-005</a>, 1961.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib69"><label>69</label><mixed-citation>
      
Tabrizi, S. E., Xiao, K., Thé, J. V. G., Saad, M., Farghaly, H., Yang, S. X., and Gharabaghi, B.: Hourly road pavement surface temperature forecasting using deep learning models, J. Hydrol., 603, 126877, <a href="https://doi.org/10.1016/j.jhydrol.2021.126877" target="_blank">https://doi.org/10.1016/j.jhydrol.2021.126877</a>, 2021.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib70"><label>70</label><mixed-citation>
      
Tao, R., Peng, R., Wang, H., Wang, J., and Qiao, J.:  Temperature and humidity prediction of mountain highway tunnel entrance road surface based on improved Bi-LSTM neural network, Evol. Syst., 15, 691–702, <a href="https://doi.org/10.1007/s12530-023-09538-5" target="_blank">https://doi.org/10.1007/s12530-023-09538-5</a>, 2024.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib71"><label>71</label><mixed-citation>
      
Vaswani, A., Shazeer, N., Parmar, N., Uszkoreit, J., Jones, L., Gomez, A. N., Kaiser, Ł., and Polosukhin, I.:  Attention is all you need, Adv. Neur. In., 30, <a href="https://doi.org/10.1201/9781003561460-19" target="_blank">https://doi.org/10.1201/9781003561460-19</a>, 2017.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib72"><label>72</label><mixed-citation>
      
Wang, D.: Simplified analytical approach to predicting asphalt pavement temperature, J. Mater. Civil Eng., 27, 04015043, <a href="https://doi.org/10.1061/(ASCE)MT.1943-5533.0001301" target="_blank">https://doi.org/10.1061/(ASCE)MT.1943-5533.0001301</a>, 2015.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib73"><label>73</label><mixed-citation>
      
Wang, D., Roesler, J. R., and Guo, D. Z.: Analytical approach to predicting temperature fields in multilayered pavement systems, J. Eng. Mech., 135, 334–344, <a href="https://doi.org/10.1061/(ASCE)0733-9399(2009)135:4(334)" target="_blank">https://doi.org/10.1061/(ASCE)0733-9399(2009)135:4(334)</a>, 2009.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib74"><label>74</label><mixed-citation>
      
Wang, X., Pan, P., and Li, J.: Real-time measurement on dynamic temperature variation of asphalt pavement using machine learning, Measurement, 207, 112413, <a href="https://doi.org/10.1016/j.measurement.2022.112413" target="_blank">https://doi.org/10.1016/j.measurement.2022.112413</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib75"><label>75</label><mixed-citation>
      
Wolpert, D. H.: Stacked generalization, Neural Netw., 5, 241–259, <a href="https://doi.org/10.1016/S0893-6080(05)80023-1" target="_blank">https://doi.org/10.1016/S0893-6080(05)80023-1</a>, 1992.


    </mixed-citation></ref-html>
<ref-html id="bib1.bib76"><label>76</label><mixed-citation>
      
Yang, C. H., Yun, D. G., Kim, J. G., Lee, G., and Kim, S. B.: Machine learning approaches to estimate road surface temperature variation along road section in real-time for winter operation, Int. J. Intell. Transp. Syst. Res., 18, 343–355, <a href="https://doi.org/10.1007/s13177-019-00198-x" target="_blank">https://doi.org/10.1007/s13177-019-00198-x</a>, 2020.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib77"><label>77</label><mixed-citation>
      
Yin, Z., Hadzimustafic, J., Kann, A., and Wang, Y.:  On statistical nowcasting of road surface temperature, Meteorol. Appl., 26, 1–13, <a href="https://doi.org/10.1002/met.1737" target="_blank">https://doi.org/10.1002/met.1737</a>, 2019.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib78"><label>78</label><mixed-citation>
      
Yuan, J., Cheng, H., Sun, L., Cao, Y., Yang, R., Jin, T., and Li, M.:  Cross-Regional Pavement Temperature Prediction Using Transfer Learning and Random Forest, Appl. Sci., 15, 7436, <a href="https://doi.org/10.3390/app15137436" target="_blank">https://doi.org/10.3390/app15137436</a>, 2025.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib79"><label>79</label><mixed-citation>
      
Zhang, M., Guo, H., Li, J. Y., Li, L., and Zhu, F.:  A deep learning approach for enhanced real-time prediction of winter road surface temperatures in high-altitude mountain areas, Promet, 36, 958–972, <a href="https://doi.org/10.7307/ptt.v36i5.541" target="_blank">https://doi.org/10.7307/ptt.v36i5.541</a>, 2024.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib80"><label>80</label><mixed-citation>
      
Zhang, N., Mao, T., Chen, H., Lv, L., Wang, Y., and Yan, Y.:  Temperature prediction for expressway pavement icing in winter based on XGBoost–LSTNet variable weight combination model, J. Transp. Eng. A-Syst., 149, 04023062, <a href="https://doi.org/10.1061/JTEPBS.TEENG-7918" target="_blank">https://doi.org/10.1061/JTEPBS.TEENG-7918</a>, 2023.

    </mixed-citation></ref-html>
<ref-html id="bib1.bib81"><label>81</label><mixed-citation>
      
Zhao, X., Shen, A., and Ma, B.: Temperature response of asphalt pavement to low temperatures and large temperature differences, Int. J. Pavement Eng., 21, 49–62, <a href="https://doi.org/10.1080/10298436.2018.1435877" target="_blank">https://doi.org/10.1080/10298436.2018.1435877</a>, 2020.

    </mixed-citation></ref-html>--></article>
