<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD JATS (Z39.96) Journal Publishing DTD v1.4 20241031//EN" "JATS-journalpublishing1-4.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article" dtd-version="1.4" xml:lang="en">
  <front>
    <journal-meta>
      <journal-id journal-id-type="publisher-id">airr</journal-id>
      <journal-title-group>
        <journal-title>Advances in Artificial Intelligence and Robotics Research</journal-title>
      </journal-title-group>
      <issn pub-type="epub">3143-3995</issn>
      <issn pub-type="ppub">3143-3987</issn>
      <publisher>
        <publisher-name>Scientific Research Publishing</publisher-name>
      </publisher>
    </journal-meta>
    <article-meta>
      <article-id pub-id-type="doi">10.4236/airr.2026.23008</article-id>
      <article-id pub-id-type="publisher-id">airr-153360</article-id>
      <article-categories>
        <subj-group>
          <subject>Article</subject>
        </subj-group>
        <subj-group>
          <subject>Computer Science</subject>
          <subject>Communications</subject>
        </subj-group>
      </article-categories>
      <title-group>
        <article-title>Diagnosing Physics-Informed Neural Networks for Motor Thermal Virtual Sensing: A Controlled Study on Single-Phase Induction Motors</article-title>
      </title-group>
      <contrib-group>
        <contrib contrib-type="author">
          <contrib-id contrib-id-type="orcid">0000-0002-1042-227X</contrib-id>
          <name name-style="western">
            <surname>Adeyi</surname>
            <given-names>Timothy A.</given-names>
          </name>
          <xref ref-type="aff" rid="aff1">1</xref>
          <xref ref-type="aff" rid="aff2">2</xref>
        </contrib>
        <contrib contrib-type="author">
          <contrib-id contrib-id-type="orcid">0000-0001-6316-2339</contrib-id>
          <name name-style="western">
            <surname>Cheok</surname>
            <given-names>Adrian D.</given-names>
          </name>
          <xref ref-type="aff" rid="aff1">1</xref>
        </contrib>
        <contrib contrib-type="author">
          <contrib-id contrib-id-type="orcid">0000-0001-8455-1915</contrib-id>
          <name name-style="western">
            <surname>Zhou</surname>
            <given-names>Steven Z.</given-names>
          </name>
          <xref ref-type="aff" rid="aff3">3</xref>
          <xref ref-type="aff" rid="aff4">4</xref>
        </contrib>
        <contrib contrib-type="author">
          <contrib-id contrib-id-type="orcid">0009-0002-3563-1788</contrib-id>
          <name name-style="western">
            <surname>Olasupo</surname>
            <given-names>Daniel O.</given-names>
          </name>
          <xref ref-type="aff" rid="aff5">5</xref>
        </contrib>
        <contrib contrib-type="author">
          <contrib-id contrib-id-type="orcid">0009-0008-0214-1356</contrib-id>
          <name name-style="western">
            <surname>Makhdoom</surname>
            <given-names>Muhammed A.</given-names>
          </name>
          <xref ref-type="aff" rid="aff6">6</xref>
        </contrib>
        <contrib contrib-type="author">
          <contrib-id contrib-id-type="orcid">0000-0001-8751-4351</contrib-id>
          <name name-style="western">
            <surname>EnochOghene</surname>
            <given-names>Samuel O.</given-names>
          </name>
          <xref ref-type="aff" rid="aff5">5</xref>
        </contrib>
      </contrib-group>
      <aff id="aff1"><label>1</label> School of Automation, Nanjing University of Information Science and Technology, Nanjing, China </aff>
      <aff id="aff2"><label>2</label> Department of Mechanical Engineering, Lead City University, Ibadan, Nigeria </aff>
      <aff id="aff3"><label>3</label> NUS Suzhou Research Institute, Suzhou, China </aff>
      <aff id="aff4"><label>4</label> Department of Electrical Computer Engineering, National University of Singapore, Singapore </aff>
      <aff id="aff5"><label>5</label> Department of Electrical &amp; Electronic Engineering, Lead City University, Ibadan, Nigeria </aff>
      <aff id="aff6"><label>6</label> School of Information &amp; Communication Engineering, Nanjing University of Information Science and Technology, Nanjing, China </aff>
      <author-notes>
        <fn fn-type="conflict" id="fn-conflict">
          <p>The authors declare no conflicts of interest regarding the publication of this paper.</p>
        </fn>
      </author-notes>
      <pub-date pub-type="epub">
        <day>01</day>
        <month>09</month>
        <year>2026</year>
      </pub-date>
      <pub-date pub-type="collection">
        <month>09</month>
        <year>2026</year>
      </pub-date>
      <volume>02</volume>
      <issue>03</issue>
      <fpage>127</fpage>
      <lpage>146</lpage>
      <history>
        <date date-type="received">
          <day>22</day>
          <month>07</month>
          <year>2026</year>
        </date>
        <date date-type="accepted">
          <day>22</day>
          <month>08</month>
          <year>2026</year>
        </date>
        <date date-type="published">
          <day>25</day>
          <month>08</month>
          <year>2026</year>
        </date>
      </history>
      <permissions>
        <copyright-statement>© 2026 by the authors and Scientific Research Publishing Inc.</copyright-statement>
        <copyright-year>2026</copyright-year>
        <license license-type="open-access">
          <license-p> This article is an open access article distributed under the terms and conditions of the Creative Commons Attribution (CC BY) license ( <ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link> ). </license-p>
        </license>
      </permissions>
      <self-uri content-type="doi" xlink:href="https://doi.org/10.4236/airr.2026.23008">https://doi.org/10.4236/airr.2026.23008</self-uri>
      <abstract>
        <p>Physics-informed neural networks (PINNs) have been proposed for motor thermal virtual sensing, yet the specific conditions under which embedded physics constraints outperform identical unconstrained networks remain unclear. This paper presents a controlled diagnostic study comparing a feedforward DNN and a PINN for winding temperature estimation in a 1 HP single-phase induction motor (SPIM). Unlike prior work, which focused on permanent magnet synchronous motors, this study addresses the distinct thermal modeling challenges of the SPIM’s dual-winding structure by incorporating a first-order thermal ODE that accounts for main/auxiliary copper losses, bearing friction, and convective cooling. We systematically evaluate performance across five standard conditions: in-distribution operation, moderate load-profile shift, severe ambient temperature shift, cold-start transients, and sensor noise. Additionally, we conduct a self-supervised experiment and a systematic lambda sensitivity analysis across seven physics loss weights. To establish boundary conditions for physics prior effectiveness, we further evaluate both models under extreme distribution shifts (E1 - E3) and degraded data scenarios (E4 - E7) encompassing sparse sampling, missing data, biased coverage, and noisy measurements. Results reveal a critical boundary condition: when temperature labels and temporal context are available, DNN and PINN achieve equivalent accuracy. However, when data are sparse, missing, biased, or noisy, the physics prior becomes valuable, reducing MAE by approximately 10% relative to DNN alone. The lambda analysis confirms that the DNN advantage is not an artifact of hyperparameter selection, while the self-supervised PINN fails (MAE &gt; 30˚C), demonstrating that the simplified first-order prior is underdetermined without label supervision. These findings establish that physics-informed regularization using a simplified first-order prior is beneficial only under degraded data conditions, providing practitioners with clear decision rules for method selection in thermal virtual sensing.</p>
      </abstract>
      <kwd-group kwd-group-type="author-generated" xml:lang="en">
        <kwd>Diagnostic Study</kwd>
        <kwd>Distribution Shift</kwd>
        <kwd>Physics-Informed Neural Networks (PINNs)</kwd>
        <kwd>Single-Phase Induction Motor (SPIM)</kwd>
        <kwd>Thermal Virtual Sensing</kwd>
      </kwd-group>
    </article-meta>
  </front>
  <body>
    <sec id="sec1">
      <title>1. Introduction</title>
      <p>Electric motors are among the most widely deployed machines in modern engineering systems. They power industrial drives, household appliances, electric vehicles, pumps, fans, and compressors across every sector of the economy [<xref ref-type="bibr" rid="B1">1</xref>]-[<xref ref-type="bibr" rid="B3">3</xref>]. In all these applications, winding temperature is a critical operating variable. Elevated temperatures accelerate insulation degradation, reduce motor efficiency, and in severe cases, lead to premature failure [<xref ref-type="bibr" rid="B4">4</xref>]. Reliable temperature monitoring is therefore essential for condition monitoring, predictive maintenance, and protection system design.</p>
      <p>Direct temperature measurement using embedded thermocouples or resistance temperature detectors is not always practical. In sealed motors, high-speed applications, and cost-sensitive consumer products, installing and maintaining physical temperature sensors introduces significant challenges in terms of cost, space, and long-term reliability [<xref ref-type="bibr" rid="B5">5</xref>]. This has motivated extensive research into virtual sensing, which involves estimating internal temperatures from electrical signals already available in the motor drive system, such as current, voltage, and speed.</p>
      <p>Two broad classes of methods have been developed for motor thermal virtual sensing. Physics-based approaches, particularly lumped-parameter thermal networks (LPTNs), model the thermal behavior of the motor using electrical analogies of heat flow. These methods provide physically guaranteed estimates and generalize well across operating conditions, but they require detailed motor geometry, multi-node calibration, and access to thermal interface resistance data that is rarely available from manufacturers [<xref ref-type="bibr" rid="B6">6</xref>][<xref ref-type="bibr" rid="B7">7</xref>]. Data-driven approaches, including support vector regression, convolutional neural networks, and long short-term memory networks, achieve high accuracy when sufficient labeled training data is available but can produce thermodynamically inconsistent predictions when tested under conditions not represented in the training set [<xref ref-type="bibr" rid="B8">8</xref>].</p>
      <p>Physics-informed neural networks (PINNs), introduced by Raissi <italic>et al</italic>. [<xref ref-type="bibr" rid="B9">9</xref>], offer a middle path between these two approaches. By embedding governing differential equations directly into the training loss function, PINNs impose physical consistency as a soft constraint during learning. Recent work has applied PINNs to motor thermal estimation with promising results. Wang <italic>et al</italic>. [<xref ref-type="bibr" rid="B10">10</xref>] combined a 10-node Motor-CAD-calibrated LPTN with neural network correction terms, achieving an MAE of 1.71˚C on real PMSM data. Stensson [<xref ref-type="bibr" rid="B11">11</xref>] demonstrated PINN-based thermal modeling for electric motors, and Lim <italic>et al</italic>. [<xref ref-type="bibr" rid="B12">12</xref>] applied physics-informed learning to electric vehicle dynamics estimation.</p>
      <p>Despite this progress, a fundamental question has not been answered: under what conditions does embedding a physics prior into the loss function actually improve performance compared to an identical unconstrained network? Existing studies propose PINN architectures and report performance on fixed test sets, but none isolates the marginal contribution of the physics constraint by holding everything else constant. If a physics prior only helps under conditions that are rarely encountered in practice, its added complexity may not be justified. If it provides robustness specifically where data-driven models fail, under distribution shift, sensor noise, or data scarcity, practitioners need to know this clearly.</p>
      <p>A second gap exists in the motor type coverage of existing work. The vast majority of PINN-based thermal estimation studies focus on permanent magnet synchronous motors [<xref ref-type="bibr" rid="B10">10</xref>][<xref ref-type="bibr" rid="B11">11</xref>][<xref ref-type="bibr" rid="B13">13</xref>]. The single-phase induction motor, which is extensively used in household appliances, small pumps, and light industrial equipment worldwide, has received almost no attention in the neural network-based thermal estimation literature. The dual-winding structure of the SPIM, with separate main and auxiliary windings, each contributing to copper losses, presents a distinct thermal modeling challenge that the PMSM-oriented literature does not address.</p>
      <p>This paper addresses both gaps through a controlled diagnostic study. A high-fidelity Simulink model of a 1 HP single-phase induction motor is used to generate five systematically varied datasets covering in-distribution operation, moderate load-profile shift, severe ambient temperature shift, cold-start transients, and noisy sensor measurements. An identical feedforward neural network architecture is trained with and without a first-order thermal ODE prior across all conditions. A self-supervised experiment, training without any temperature labels, is also conducted to establish the standalone capability of the physics prior.</p>
      <p>The contributions of this paper are as follows:</p>
      <p>1) The first controlled diagnostic comparison of DNN and PINN for single-phase induction motor thermal estimation uses identical architectures, varying only the loss function.</p>
      <p>2) A dual-winding thermal ODE formulation for SPIM that explicitly accounts for main and auxiliary winding copper losses and bearing friction losses, following Calasan <italic>et al</italic>. [<xref ref-type="bibr" rid="B14">14</xref>].</p>
      <p>3) A systematic lambda sensitivity analysis demonstrates that the DNN advantage over the PINN is not an artifact of hyperparameter selection but a consistent property across the full range of physics loss weights tested.</p>
      <p>4) A self-supervised experiment to establish the minimum data requirements for physics-informed thermal estimation.</p>
      <p><bold>Data and Code Availability:</bold> To ensure reproducibility and facilitate further research, the Simulink model, Python training code, generated datasets, and all figures used in this study are publicly available at <ext-link ext-link-type="uri" xlink:href="https://github.com/Pstone123/SPIM-Thermal-PINN-v2">https://github.com/Pstone123/SPIM-Thermal-PINN-v2</ext-link>.</p>
    </sec>
    <sec id="sec2">
      <title>2. Related Work</title>
      <sec id="sec2dot1">
        <title>2.1. Motor Thermal Estimation</title>
        <p>The thermal monitoring of electric motors has been comprehensively reviewed by Wallscheid [<xref ref-type="bibr" rid="B5">5</xref>], who identified virtual sensing as a critical direction for condition monitoring in sealed and high-speed applications. Physics-based approaches have been the traditional standard. Wallscheid and Bocker [<xref ref-type="bibr" rid="B6">6</xref>] developed a globally identified four-node LPTN for PMSMs, demonstrating that accurate thermal modeling can be achieved with relatively few parameters when the network topology is well chosen. Kral <italic>et al</italic>. [<xref ref-type="bibr" rid="B7">7</xref>] proposed a practical two-node model for estimating permanent magnet and stator winding temperatures, which has since been widely adopted as a benchmark for lumped-parameter approaches. Chowdhury and Baski [<xref ref-type="bibr" rid="B15">15</xref>] demonstrated a simple lumped-parameter thermal model for totally enclosed fan-cooled machines, highlighting the trade-off between model simplicity and accuracy.</p>
        <p>Data-driven approaches have advanced significantly over the past decade. The Paderborn PMSM benchmark introduced by Kirchgässner <italic>et al</italic>. [<xref ref-type="bibr" rid="B16">16</xref>] provided the research community with a publicly available real test bench dataset covering 139 hours of automotive PMSM operation, enabling reproducible comparison of data-driven methods. Jin <italic>et al</italic>. [<xref ref-type="bibr" rid="B8">8</xref>] proposed a model-based and data-driven integrated temperature estimation method that combines a simplified LPTN with a neural network correction layer, achieving strong performance on real motor data. These hybrid approaches represent an intermediate position between pure physics-based and pure data-driven methods.</p>
      </sec>
      <sec id="sec2dot2">
        <title>2.2. Physics-Informed Neural Networks</title>
        <p>The PINN framework, introduced by Raissi <italic>et al</italic>. [<xref ref-type="bibr" rid="B9">9</xref>], embeds governing differential equations into the loss function through an additional residual term computed via automatic differentiation. This approach has been applied across a wide range of engineering domains, from fluid dynamics to structural mechanics. In the motor thermal estimation domain, Wang <italic>et al</italic>. [<xref ref-type="bibr" rid="B10">10</xref>] demonstrated that combining a high-fidelity 10-node LPTN prior with neural network correction terms achieves an MAE of 1.71˚C on real PMSM data, substantially better than prior purely data-driven results on the same hardware. Stensson [<xref ref-type="bibr" rid="B11">11</xref>] showed that PINN-based thermal modeling can be applied to electric motors using supervised training with temperature labels. Lim <italic>et al</italic>. [<xref ref-type="bibr" rid="B12">12</xref>] applied physics-informed learning to electric vehicle dynamics, demonstrating the potential of minimal-input PINN models for automotive applications.</p>
        <p>Reinbold <italic>et al</italic>. [<xref ref-type="bibr" rid="B17">17</xref>] demonstrated that physically constrained learning can improve robustness on noisy, incomplete experimental data, which motivated the noise robustness experiment in the present study. Chen <italic>et al</italic>. [<xref ref-type="bibr" rid="B18">18</xref>] proposed gradient normalization as a technique for balancing competing loss terms in multi-task neural networks, which is directly relevant to the challenge of scaling the physics and data loss terms in PINN training, a challenge explicitly addressed in this paper through residual normalization.</p>
      </sec>
      <sec id="sec2dot3">
        <title>2.3. Diagnostic Studies of PINN Performance</title>
        <p>The broader PINN literature has begun to examine conditions under which physics-informed learning fails or underperforms. Krishnapriyan <italic>et al</italic>. [<xref ref-type="bibr" rid="B19">19</xref>] characterized failure modes in PINNs applied to stiff partial differential equations, identifying gradient pathologies and loss landscape properties that prevent convergence. These studies focus on forward and inverse PDE problems in computational physics rather than on engineering state estimation where the governing equation is a simplified approximation of the true system dynamics. No equivalent diagnostic study exists for motor thermal estimation. The present work fills this gap by providing the first systematic comparison of identical DNN and PINN architectures across controlled variations of data condition, distribution shift type, and physics loss weighting for motor thermal estimation.</p>
      </sec>
    </sec>
    <sec id="sec3">
      <title>3. Methodology</title>
      <sec id="sec3dot1">
        <title>3.1. SPIM Simulation Model</title>
        <p>A single-phase induction motor simulation model was developed in MATLAB/ Simulink R2025b using the Simscape Electrical capacitor-start-run block. The motor was parameterized using manufacturer datasheet values for a 1 HP, 230 V, 50 Hz SPIM. The main winding stator resistance is 2.02 Ω, the auxiliary winding stator resistance is 7.14 Ω, and the rotor resistance referred to the stator is 4.12 Ω. The rotor inertia is 0.0146 kg·m<sup>2</sup>. The motor was supplied through an H-bridge inverter configuration, and the mechanical load applied at the shaft port was controlled through a time-varying MATLAB function block.</p>
        <p>The mechanical load torque was designed as a stepped time-varying profile to create rich, dynamic operating conditions suitable for neural network training. For the training dataset, the load torque alternated between 2.0 Nm and 3.5 Nm at 20-second intervals over a total simulation duration of 200 seconds. This represents a realistic scenario in which a motor experiences varying mechanical demand, such as a pump or fan operating under changing flow conditions. The load values were selected to maintain motor operation within the normal torque range, with rated torque estimated at approximately 4.75 Nm for a 1 HP, 50 Hz, four-pole machine.</p>
      </sec>
      <sec id="sec3dot2">
        <title>3.2. Thermal Subsystem</title>
        <p>A first-order lumped thermal subsystem was coupled to the electrical model to compute winding temperature dynamics. The governing thermal equation is:</p>
        <disp-formula id="FD1">
          <label>(1)</label>
          <mml:math>
            <mml:mrow>
              <mml:msub>
                <mml:mi>C</mml:mi>
                <mml:mrow>
                  <mml:mi>t</mml:mi>
                  <mml:mi>h</mml:mi>
                </mml:mrow>
              </mml:msub>
              <mml:mfrac>
                <mml:mrow>
                  <mml:mtext>d</mml:mtext>
                  <mml:mi>T</mml:mi>
                </mml:mrow>
                <mml:mrow>
                  <mml:mtext>d</mml:mtext>
                  <mml:mi>t</mml:mi>
                </mml:mrow>
              </mml:mfrac>
              <mml:mo>=</mml:mo>
              <mml:msubsup>
                <mml:mi>I</mml:mi>
                <mml:mrow>
                  <mml:mi>m</mml:mi>
                  <mml:mi>a</mml:mi>
                  <mml:mi>i</mml:mi>
                  <mml:mi>n</mml:mi>
                </mml:mrow>
                <mml:mn>2</mml:mn>
              </mml:msubsup>
              <mml:msub>
                <mml:mi>R</mml:mi>
                <mml:mrow>
                  <mml:mi>m</mml:mi>
                  <mml:mi>a</mml:mi>
                  <mml:mi>i</mml:mi>
                  <mml:mi>n</mml:mi>
                </mml:mrow>
              </mml:msub>
              <mml:mrow>
                <mml:mo>(</mml:mo>
                <mml:mi>T</mml:mi>
                <mml:mo>)</mml:mo>
              </mml:mrow>
              <mml:mo>+</mml:mo>
              <mml:msubsup>
                <mml:mi>I</mml:mi>
                <mml:mrow>
                  <mml:mi>a</mml:mi>
                  <mml:mi>u</mml:mi>
                  <mml:mi>x</mml:mi>
                </mml:mrow>
                <mml:mn>2</mml:mn>
              </mml:msubsup>
              <mml:msub>
                <mml:mi>R</mml:mi>
                <mml:mrow>
                  <mml:mi>a</mml:mi>
                  <mml:mi>u</mml:mi>
                  <mml:mi>x</mml:mi>
                </mml:mrow>
              </mml:msub>
              <mml:mrow>
                <mml:mo>(</mml:mo>
                <mml:mi>T</mml:mi>
                <mml:mo>)</mml:mo>
              </mml:mrow>
              <mml:mo>+</mml:mo>
              <mml:mi>B</mml:mi>
              <mml:mi>ω</mml:mi>
              <mml:mo>−</mml:mo>
              <mml:mi>h</mml:mi>
              <mml:mrow>
                <mml:mo>(</mml:mo>
                <mml:mrow>
                  <mml:mi>T</mml:mi>
                  <mml:mo>−</mml:mo>
                  <mml:msub>
                    <mml:mi>T</mml:mi>
                    <mml:mrow>
                      <mml:mi>a</mml:mi>
                      <mml:mi>m</mml:mi>
                      <mml:mi>b</mml:mi>
                    </mml:mrow>
                  </mml:msub>
                </mml:mrow>
                <mml:mo>)</mml:mo>
              </mml:mrow>
            </mml:mrow>
          </mml:math>
        </disp-formula>
        <p>where <italic>T</italic> is the winding temperature in degrees Celsius, <italic>C</italic><italic><sub>th</sub></italic> = 8000 J/˚C is the thermal capacitance, <italic>h</italic> = 28 W/˚C is the convective heat transfer coefficient, <italic>T</italic><italic><sub>amb</sub></italic> is the ambient temperature, and <italic>ω</italic> is the rotor angular speed in rad/s. The first two terms represent copper losses in the main and auxiliary windings, respectively. The temperature dependence of copper resistance is modeled as:</p>
        <disp-formula id="FD2">
          <label>(2)</label>
          <mml:math display="inline">
            <mml:mrow>
              <mml:mi>R</mml:mi>
              <mml:mrow>
                <mml:mo>(</mml:mo>
                <mml:mi>T</mml:mi>
                <mml:mo>)</mml:mo>
              </mml:mrow>
              <mml:mo>=</mml:mo>
              <mml:msub>
                <mml:mi>R</mml:mi>
                <mml:mn>0</mml:mn>
              </mml:msub>
              <mml:mrow>
                <mml:mo>[</mml:mo>
                <mml:mrow>
                  <mml:mn>1</mml:mn>
                  <mml:mo>+</mml:mo>
                  <mml:mi>α</mml:mi>
                  <mml:mrow>
                    <mml:mo>(</mml:mo>
                    <mml:mrow>
                      <mml:mi>T</mml:mi>
                      <mml:mo>−</mml:mo>
                      <mml:msub>
                        <mml:mi>T</mml:mi>
                        <mml:mrow>
                          <mml:mi>r</mml:mi>
                          <mml:mi>e</mml:mi>
                          <mml:mi>f</mml:mi>
                        </mml:mrow>
                      </mml:msub>
                    </mml:mrow>
                    <mml:mo>)</mml:mo>
                  </mml:mrow>
                </mml:mrow>
                <mml:mo>]</mml:mo>
              </mml:mrow>
            </mml:mrow>
          </mml:math>
        </disp-formula>
        <p>where <italic>R</italic><sub>0</sub> is the reference resistance at <italic>T</italic><italic><sub>ref</sub></italic>, <italic>α</italic> = 0.00393/˚C is the temperature coefficient of copper, and <italic>T</italic><italic><sub>ref</sub></italic> = 25˚C is the reference temperature. The third term, <italic>B</italic>·<italic>ω</italic>, represents friction and bearing losses following the viscous friction model proposed by Calasan <italic>et al</italic>. [<xref ref-type="bibr" rid="B14">14</xref>], where <italic>B</italic> = 0.0005 Nm·s/rad is the bearing loss coefficient. The fourth term represents convective cooling to the ambient environment. Rotor copper losses were not included, as rotor current is not directly observable from stator measurements in this configuration. Core iron losses were excluded, as the iron loss coefficient is not available from the manufacturer datasheet. Both exclusions are consistent with virtual sensing approaches in the literature and are acknowledged as simplifying assumptions.</p>
        <p>The model was solved using the variable-step ode45 solver with automatic step size control. Simulation outputs were sampled at 10 Hz (sample interval Δ<italic>t</italic> = 0.1 s), yielding 2001 samples per 200-second simulation. Current signals were processed through RMS blocks with a 0.1-second sliding window prior to logging, converting raw AC waveforms to effective heating values appropriate for thermal modeling.</p>
      </sec>
      <sec id="sec3dot3">
        <title>3.3. Dataset Generation</title>
        <p>Five datasets were generated by systematically varying the mechanical load profile and ambient temperature conditions. All datasets share the same input feature set: rotor speed (RPM), main winding RMS current (A), auxiliary winding RMS current (A), electromagnetic torque (Nm), elapsed time (s), and ambient temperature (˚C). The target variable is winding temperature (˚C). <bold>Table 1</bold> summarizes the dataset properties.</p>
        <p>Dataset A served as the sole training source for all models. Datasets B1, B2, C, and D were used exclusively for evaluation; no model was retrained or fine-tuned on these sets. This design ensures that reported performance differences reflect genuine generalization capability rather than memorization.</p>
        <p>For Dataset D, additive Gaussian noise with a standard deviation equal to 5% of each signal’s full-scale range was applied independently to speed, main current, and auxiliary current measurements after simulation. Temperature labels remained clean, reflecting the assumption that temperature sensors are more reliable than electrical measurement channels. The resulting temperature trajectories for these datasets are visualized in <xref ref-type="fig" rid="fig1">Figure 1</xref>.</p>
        <p><bold>Table 1</bold><bold>.</bold> SPIM dataset summary.</p>
        <table-wrap id="tbl1">
          <label>Table 1</label>
          <table>
            <tbody>
              <tr>
                <td>Dataset</td>
                <td>Purpose</td>
                <td>Load Profile</td>
                <td>Ambient (˚C)</td>
                <td>Duration (s)</td>
                <td>Samples</td>
              </tr>
              <tr>
                <td>A</td>
                <td>Training</td>
                <td>Stepped 2.0 - 3.5 Nm, 20 s intervals</td>
                <td>25</td>
                <td>200</td>
                <td>2001</td>
              </tr>
              <tr>
                <td>B1</td>
                <td>Moderate shift</td>
                <td>Different step patterns, varied magnitudes</td>
                <td>25</td>
                <td>200</td>
                <td>2001</td>
              </tr>
              <tr>
                <td>B2</td>
                <td>Severe shift</td>
                <td>Same as A, elevated ambient.</td>
                <td>40</td>
                <td>200</td>
                <td>2001</td>
              </tr>
              <tr>
                <td>C</td>
                <td>Cold start</td>
                <td>Same as A, early transient only.</td>
                <td>25</td>
                <td>100</td>
                <td>1001</td>
              </tr>
              <tr>
                <td>D</td>
                <td>Sensor noise</td>
                <td>Same as A, 5% Gaussian noise on inputs.</td>
                <td>25</td>
                <td>200</td>
                <td>2001</td>
              </tr>
            </tbody>
          </table>
        </table-wrap>
        <fig id="fig1">
          <label>Figure 1</label>
          <graphic xlink:href="https://html.scirp.org/file/2970142-rId25.jpeg?20260825114555" />
        </fig>
        <p><bold>Figure 1</bold><bold>.</bold> Winding temperature trajectories for all four SPIM datasets.</p>
        <p>To systematically investigate the boundary conditions of physics prior effectiveness, seven additional datasets (E1 - E7) were generated. Datasets E1 - E3 represent extreme distribution shifts: E1 (50˚C ambient temperature), E2 (2× load scaling), and E3 (extended 500 s duration). Datasets E4 - E7 represent degraded data scenarios: E4 (sparse sampling at 0.5 s intervals), E5 (missing data with 20% of samples randomly removed), E6 (biased coverage with training restricted to high-torque operating regions), and E7 (noisy inputs with 10% Gaussian noise). These datasets are used exclusively for evaluation; no model is retrained on them.</p>
      </sec>
      <sec id="sec3dot4">
        <title>3.4. Neural Network Architecture</title>
        <p>Both the DNN and PINN share an identical feedforward architecture to ensure a fair, controlled comparison. The network takes five inputs: <italic>ω</italic>, <italic>I</italic><italic><sub>main</sub></italic>, <italic>I</italic><italic><sub>aux</sub></italic>, <italic>T</italic><italic><sub>e</sub></italic>, and <italic>t</italic>, which are rotor speed, main winding current, auxiliary winding current, torque, and time, respectively, with a single temperature output (<italic>T</italic>). Three hidden layers with 20, 20, and 10 neurons, respectively, use hyperbolic tangent (tanh) activation functions. The output layer is linear. This gives 761 trainable parameters, a deliberately compact architecture chosen to avoid overfitting on the relatively small training set of 2001 samples.</p>
        <p>All inputs and the temperature output were normalized to [0, 1] using the MinMaxScaler fitted exclusively on Dataset A. The same scaler was applied to all other datasets without refitting, preserving the distributional mismatch that defines the shift conditions.</p>
      </sec>
      <sec id="sec3dot5">
        <title>3.5. Loss Functions</title>
        <p>The DNN was trained using the mean squared error between predicted and true normalized temperature:</p>
        <disp-formula id="FD3">
          <label>(3)</label>
          <mml:math>
            <mml:mrow>
              <mml:msub>
                <mml:mi>L</mml:mi>
                <mml:mrow>
                  <mml:mi>D</mml:mi>
                  <mml:mi>N</mml:mi>
                  <mml:mi>N</mml:mi>
                </mml:mrow>
              </mml:msub>
              <mml:mo>=</mml:mo>
              <mml:mfrac>
                <mml:mn>1</mml:mn>
                <mml:mi>N</mml:mi>
              </mml:mfrac>
              <mml:mstyle displaystyle="true">
                <mml:msubsup>
                  <mml:mo>∑</mml:mo>
                  <mml:mrow>
                    <mml:mi>i</mml:mi>
                    <mml:mo>=</mml:mo>
                    <mml:mn>1</mml:mn>
                  </mml:mrow>
                  <mml:mi>N</mml:mi>
                </mml:msubsup>
                <mml:mrow>
                  <mml:msup>
                    <mml:mrow>
                      <mml:mrow>
                        <mml:mo>(</mml:mo>
                        <mml:mrow>
                          <mml:msub>
                            <mml:mover accent="true">
                              <mml:mi>T</mml:mi>
                              <mml:mo>^</mml:mo>
                            </mml:mover>
                            <mml:mi>i</mml:mi>
                          </mml:msub>
                          <mml:mo>−</mml:mo>
                          <mml:msub>
                            <mml:mi>T</mml:mi>
                            <mml:mrow>
                              <mml:mi>t</mml:mi>
                              <mml:mi>r</mml:mi>
                              <mml:mi>u</mml:mi>
                              <mml:mi>e</mml:mi>
                              <mml:mo>,</mml:mo>
                              <mml:mi>i</mml:mi>
                            </mml:mrow>
                          </mml:msub>
                        </mml:mrow>
                        <mml:mo>)</mml:mo>
                      </mml:mrow>
                    </mml:mrow>
                    <mml:mn>2</mml:mn>
                  </mml:msup>
                </mml:mrow>
              </mml:mstyle>
            </mml:mrow>
          </mml:math>
        </disp-formula>
        <p>The PINN was trained using a combined loss:</p>
        <disp-formula id="FD4">
          <label>(4)</label>
          <mml:math display="inline">
            <mml:mrow>
              <mml:msub>
                <mml:mi>L</mml:mi>
                <mml:mrow>
                  <mml:mi>P</mml:mi>
                  <mml:mi>I</mml:mi>
                  <mml:mi>N</mml:mi>
                  <mml:mi>N</mml:mi>
                </mml:mrow>
              </mml:msub>
              <mml:mo>=</mml:mo>
              <mml:msub>
                <mml:mi>L</mml:mi>
                <mml:mrow>
                  <mml:mi>d</mml:mi>
                  <mml:mi>a</mml:mi>
                  <mml:mi>t</mml:mi>
                  <mml:mi>a</mml:mi>
                </mml:mrow>
              </mml:msub>
              <mml:mo>+</mml:mo>
              <mml:mi>λ</mml:mi>
              <mml:mo>⋅</mml:mo>
              <mml:msub>
                <mml:mi>L</mml:mi>
                <mml:mrow>
                  <mml:mi>p</mml:mi>
                  <mml:mi>h</mml:mi>
                  <mml:mi>y</mml:mi>
                  <mml:mi>s</mml:mi>
                  <mml:mi>i</mml:mi>
                  <mml:mi>c</mml:mi>
                  <mml:mi>s</mml:mi>
                </mml:mrow>
              </mml:msub>
            </mml:mrow>
          </mml:math>
        </disp-formula>
        <p>where <italic>λ</italic> is the physics loss weight and <italic>L</italic><italic><sub>physics</sub></italic> is the mean squared residual of the thermal ODE evaluated at each training point. The time derivative (<inline-formula><mml:math><mml:mrow><mml:mfrac><mml:mrow><mml:mtext> d </mml:mtext><mml:mover accent="true"><mml:mi> T </mml:mi><mml:mo> ^ </mml:mo></mml:mover></mml:mrow><mml:mrow><mml:mtext> d </mml:mtext><mml:mi> t </mml:mi></mml:mrow></mml:mfrac></mml:mrow></mml:math></inline-formula>) was computed via automatic differentiation with respect to the time input, following the standard PINN formulation of Raissi <italic>et al</italic>. [<xref ref-type="bibr" rid="B9">9</xref>]. The physics residual is:</p>
        <disp-formula id="FD5">
          <label>(5)</label>
          <mml:math display="inline">
            <mml:mrow>
              <mml:msub>
                <mml:mi>L</mml:mi>
                <mml:mrow>
                  <mml:mi>p</mml:mi>
                  <mml:mi>h</mml:mi>
                  <mml:mi>y</mml:mi>
                  <mml:mi>s</mml:mi>
                  <mml:mi>i</mml:mi>
                  <mml:mi>c</mml:mi>
                  <mml:mi>s</mml:mi>
                </mml:mrow>
              </mml:msub>
              <mml:mo>=</mml:mo>
              <mml:mfrac>
                <mml:mn>1</mml:mn>
                <mml:mi>N</mml:mi>
              </mml:mfrac>
              <mml:mstyle displaystyle="true">
                <mml:msubsup>
                  <mml:mo>∑</mml:mo>
                  <mml:mrow>
                    <mml:mi>i</mml:mi>
                    <mml:mo>=</mml:mo>
                    <mml:mn>1</mml:mn>
                  </mml:mrow>
                  <mml:mi>N</mml:mi>
                </mml:msubsup>
                <mml:mrow>
                  <mml:msup>
                    <mml:mrow>
                      <mml:mrow>
                        <mml:mo>[</mml:mo>
                        <mml:mrow>
                          <mml:mfrac>
                            <mml:mrow>
                              <mml:msub>
                                <mml:mi>C</mml:mi>
                                <mml:mrow>
                                  <mml:mi>t</mml:mi>
                                  <mml:mi>h</mml:mi>
                                </mml:mrow>
                              </mml:msub>
                              <mml:mfrac>
                                <mml:mrow>
                                  <mml:mtext>d</mml:mtext>
                                  <mml:mover accent="true">
                                    <mml:mi>T</mml:mi>
                                    <mml:mo>^</mml:mo>
                                  </mml:mover>
                                </mml:mrow>
                                <mml:mrow>
                                  <mml:mtext>d</mml:mtext>
                                  <mml:mi>t</mml:mi>
                                </mml:mrow>
                              </mml:mfrac>
                              <mml:mo>−</mml:mo>
                              <mml:mrow>
                                <mml:mo>(</mml:mo>
                                <mml:mrow>
                                  <mml:msubsup>
                                    <mml:mi>I</mml:mi>
                                    <mml:mrow>
                                      <mml:mi>m</mml:mi>
                                      <mml:mi>a</mml:mi>
                                      <mml:mi>i</mml:mi>
                                      <mml:mi>n</mml:mi>
                                    </mml:mrow>
                                    <mml:mn>2</mml:mn>
                                  </mml:msubsup>
                                  <mml:msub>
                                    <mml:mi>R</mml:mi>
                                    <mml:mrow>
                                      <mml:mi>m</mml:mi>
                                      <mml:mi>a</mml:mi>
                                      <mml:mi>i</mml:mi>
                                      <mml:mi>n</mml:mi>
                                    </mml:mrow>
                                  </mml:msub>
                                  <mml:mrow>
                                    <mml:mo>(</mml:mo>
                                    <mml:mover accent="true">
                                      <mml:mi>T</mml:mi>
                                      <mml:mo>^</mml:mo>
                                    </mml:mover>
                                    <mml:mo>)</mml:mo>
                                  </mml:mrow>
                                  <mml:mo>+</mml:mo>
                                  <mml:msubsup>
                                    <mml:mi>I</mml:mi>
                                    <mml:mrow>
                                      <mml:mi>a</mml:mi>
                                      <mml:mi>u</mml:mi>
                                      <mml:mi>x</mml:mi>
                                    </mml:mrow>
                                    <mml:mn>2</mml:mn>
                                  </mml:msubsup>
                                  <mml:msub>
                                    <mml:mi>R</mml:mi>
                                    <mml:mrow>
                                      <mml:mi>a</mml:mi>
                                      <mml:mi>u</mml:mi>
                                      <mml:mi>x</mml:mi>
                                    </mml:mrow>
                                  </mml:msub>
                                  <mml:mrow>
                                    <mml:mo>(</mml:mo>
                                    <mml:mover accent="true">
                                      <mml:mi>T</mml:mi>
                                      <mml:mo>^</mml:mo>
                                    </mml:mover>
                                    <mml:mo>)</mml:mo>
                                  </mml:mrow>
                                  <mml:mo>+</mml:mo>
                                  <mml:mi>B</mml:mi>
                                  <mml:mi>ω</mml:mi>
                                  <mml:mo>−</mml:mo>
                                  <mml:mi>h</mml:mi>
                                  <mml:mrow>
                                    <mml:mo>(</mml:mo>
                                    <mml:mrow>
                                      <mml:mover accent="true">
                                        <mml:mi>T</mml:mi>
                                        <mml:mo>^</mml:mo>
                                      </mml:mover>
                                      <mml:mo>−</mml:mo>
                                      <mml:msub>
                                        <mml:mi>T</mml:mi>
                                        <mml:mrow>
                                          <mml:mi>a</mml:mi>
                                          <mml:mi>m</mml:mi>
                                          <mml:mi>b</mml:mi>
                                        </mml:mrow>
                                      </mml:msub>
                                    </mml:mrow>
                                    <mml:mo>)</mml:mo>
                                  </mml:mrow>
                                </mml:mrow>
                                <mml:mo>)</mml:mo>
                              </mml:mrow>
                            </mml:mrow>
                            <mml:mrow>
                              <mml:mn>1000</mml:mn>
                            </mml:mrow>
                          </mml:mfrac>
                        </mml:mrow>
                        <mml:mo>]</mml:mo>
                      </mml:mrow>
                    </mml:mrow>
                    <mml:mn>2</mml:mn>
                  </mml:msup>
                </mml:mrow>
              </mml:mstyle>
            </mml:mrow>
          </mml:math>
        </disp-formula>
        <p>The division by 1000<sup>2</sup> normalizes the residual, which operates in units of watts squared, to the same order of magnitude as the data loss, which operates in normalized temperature units. Without this scaling, the physics gradient overwhelms the data gradient at initialization, preventing convergence to a physically and empirically consistent solution.</p>
        <p>The self-supervised PINN was trained with <italic>λ</italic> → ∞ (<italic>L</italic><italic><sub>data</sub></italic> = 0), using only the physics residual as the training signal. This establishes the standalone estimation capability of the thermal prior without any temperature label supervision.</p>
      </sec>
      <sec id="sec3dot6">
        <title>3.6. Training Protocol</title>
        <p>All models were implemented in PyTorch and trained for 3000 epochs using the Adam optimizer with a fixed learning rate of 1 × 10<sup>−</sup><sup>3</sup> and full-batch gradient updates. A fixed random seed (42) was used throughout for reproducibility. Training was performed on the CPU using Google Colab.</p>
        <p>A systematic lambda sensitivity analysis was performed by training separate PINN instances with <italic>λ</italic> ∈ {1 × 10<sup>−</sup><sup>6</sup>, 5 × 10<sup>−</sup><sup>6</sup>, 1 × 10<sup>−</sup><sup>5</sup>, 5 × 10<sup>−</sup><sup>5</sup>, 1 × 10<sup>−</sup><sup>4</sup>, 5 × 10<sup>−</sup><sup>4</sup>, 1 × 10<sup>−</sup><sup>3</sup>} and evaluating each on both Dataset A and Dataset B1. This analysis establishes whether the DNN superiority observed at any single <italic>λ</italic> value is an artifact of hyperparameter selection or a consistent property across the physics loss weighting range.</p>
      </sec>
      <sec id="sec3dot7">
        <title>3.7. Evaluation Metrics</title>
        <p>Model performance was quantified using three standard metrics for regression tasks in motor thermal estimation:</p>
        <disp-formula id="FD6">
          <label>(6)</label>
          <mml:math>
            <mml:mrow>
              <mml:mtext>MAE</mml:mtext>
              <mml:mo>=</mml:mo>
              <mml:mfrac>
                <mml:mn>1</mml:mn>
                <mml:mi>N</mml:mi>
              </mml:mfrac>
              <mml:mstyle displaystyle="true">
                <mml:msubsup>
                  <mml:mo>∑</mml:mo>
                  <mml:mrow>
                    <mml:mi>i</mml:mi>
                    <mml:mo>=</mml:mo>
                    <mml:mn>1</mml:mn>
                  </mml:mrow>
                  <mml:mi>N</mml:mi>
                </mml:msubsup>
                <mml:mrow>
                  <mml:mrow>
                    <mml:mo>|</mml:mo>
                    <mml:mrow>
                      <mml:msub>
                        <mml:mover accent="true">
                          <mml:mi>T</mml:mi>
                          <mml:mo>^</mml:mo>
                        </mml:mover>
                        <mml:mi>i</mml:mi>
                      </mml:msub>
                      <mml:mo>−</mml:mo>
                      <mml:msub>
                        <mml:mi>T</mml:mi>
                        <mml:mrow>
                          <mml:mi>t</mml:mi>
                          <mml:mi>r</mml:mi>
                          <mml:mi>u</mml:mi>
                          <mml:mi>e</mml:mi>
                          <mml:mo>,</mml:mo>
                          <mml:mi>i</mml:mi>
                        </mml:mrow>
                      </mml:msub>
                    </mml:mrow>
                    <mml:mo>|</mml:mo>
                  </mml:mrow>
                </mml:mrow>
              </mml:mstyle>
            </mml:mrow>
          </mml:math>
        </disp-formula>
        <disp-formula id="FD7">
          <label>(7)</label>
          <mml:math>
            <mml:mrow>
              <mml:mtext>RMSE</mml:mtext>
              <mml:mo>=</mml:mo>
              <mml:msqrt>
                <mml:mrow>
                  <mml:mfrac>
                    <mml:mn>1</mml:mn>
                    <mml:mi>N</mml:mi>
                  </mml:mfrac>
                  <mml:mstyle displaystyle="true">
                    <mml:msubsup>
                      <mml:mo>∑</mml:mo>
                      <mml:mrow>
                        <mml:mi>i</mml:mi>
                        <mml:mo>=</mml:mo>
                        <mml:mn>1</mml:mn>
                      </mml:mrow>
                      <mml:mi>N</mml:mi>
                    </mml:msubsup>
                    <mml:mrow>
                      <mml:msup>
                        <mml:mrow>
                          <mml:mrow>
                            <mml:mo>(</mml:mo>
                            <mml:mrow>
                              <mml:msub>
                                <mml:mover accent="true">
                                  <mml:mi>T</mml:mi>
                                  <mml:mo>^</mml:mo>
                                </mml:mover>
                                <mml:mi>i</mml:mi>
                              </mml:msub>
                              <mml:mo>−</mml:mo>
                              <mml:msub>
                                <mml:mi>T</mml:mi>
                                <mml:mrow>
                                  <mml:mi>t</mml:mi>
                                  <mml:mi>r</mml:mi>
                                  <mml:mi>u</mml:mi>
                                  <mml:mi>e</mml:mi>
                                  <mml:mo>,</mml:mo>
                                  <mml:mi>i</mml:mi>
                                </mml:mrow>
                              </mml:msub>
                            </mml:mrow>
                            <mml:mo>)</mml:mo>
                          </mml:mrow>
                        </mml:mrow>
                        <mml:mn>2</mml:mn>
                      </mml:msup>
                    </mml:mrow>
                  </mml:mstyle>
                </mml:mrow>
              </mml:msqrt>
            </mml:mrow>
          </mml:math>
        </disp-formula>
        <disp-formula id="FD8">
          <label>(8)</label>
          <mml:math>
            <mml:mrow>
              <mml:msup>
                <mml:mtext>R</mml:mtext>
                <mml:mn>2</mml:mn>
              </mml:msup>
              <mml:mo>=</mml:mo>
              <mml:mn>1</mml:mn>
              <mml:mo>−</mml:mo>
              <mml:mfrac>
                <mml:mrow>
                  <mml:mstyle displaystyle="true">
                    <mml:msubsup>
                      <mml:mo>∑</mml:mo>
                      <mml:mrow>
                        <mml:mi>i</mml:mi>
                        <mml:mo>=</mml:mo>
                        <mml:mn>1</mml:mn>
                      </mml:mrow>
                      <mml:mi>N</mml:mi>
                    </mml:msubsup>
                    <mml:mrow>
                      <mml:msup>
                        <mml:mrow>
                          <mml:mrow>
                            <mml:mo>(</mml:mo>
                            <mml:mrow>
                              <mml:msub>
                                <mml:mi>T</mml:mi>
                                <mml:mrow>
                                  <mml:mi>t</mml:mi>
                                  <mml:mi>r</mml:mi>
                                  <mml:mi>u</mml:mi>
                                  <mml:mi>e</mml:mi>
                                  <mml:mo>,</mml:mo>
                                  <mml:mi>i</mml:mi>
                                </mml:mrow>
                              </mml:msub>
                              <mml:mo>−</mml:mo>
                              <mml:msub>
                                <mml:mover accent="true">
                                  <mml:mi>T</mml:mi>
                                  <mml:mo>^</mml:mo>
                                </mml:mover>
                                <mml:mi>i</mml:mi>
                              </mml:msub>
                            </mml:mrow>
                            <mml:mo>)</mml:mo>
                          </mml:mrow>
                        </mml:mrow>
                        <mml:mn>2</mml:mn>
                      </mml:msup>
                    </mml:mrow>
                  </mml:mstyle>
                </mml:mrow>
                <mml:mrow>
                  <mml:mstyle displaystyle="true">
                    <mml:msubsup>
                      <mml:mo>∑</mml:mo>
                      <mml:mrow>
                        <mml:mi>i</mml:mi>
                        <mml:mo>=</mml:mo>
                        <mml:mn>1</mml:mn>
                      </mml:mrow>
                      <mml:mi>N</mml:mi>
                    </mml:msubsup>
                    <mml:mrow>
                      <mml:msup>
                        <mml:mrow>
                          <mml:mrow>
                            <mml:mo>(</mml:mo>
                            <mml:mrow>
                              <mml:msub>
                                <mml:mi>T</mml:mi>
                                <mml:mrow>
                                  <mml:mi>t</mml:mi>
                                  <mml:mi>r</mml:mi>
                                  <mml:mi>u</mml:mi>
                                  <mml:mi>e</mml:mi>
                                  <mml:mo>,</mml:mo>
                                  <mml:mi>i</mml:mi>
                                </mml:mrow>
                              </mml:msub>
                              <mml:mo>−</mml:mo>
                              <mml:msub>
                                <mml:mover accent="true">
                                  <mml:mi>T</mml:mi>
                                  <mml:mo>¯</mml:mo>
                                </mml:mover>
                                <mml:mi>i</mml:mi>
                              </mml:msub>
                            </mml:mrow>
                            <mml:mo>)</mml:mo>
                          </mml:mrow>
                        </mml:mrow>
                        <mml:mn>2</mml:mn>
                      </mml:msup>
                    </mml:mrow>
                  </mml:mstyle>
                </mml:mrow>
              </mml:mfrac>
            </mml:mrow>
          </mml:math>
        </disp-formula>
        <p>where <inline-formula><mml:math><mml:mover accent="true"><mml:mi> T </mml:mi><mml:mo> ¯ </mml:mo></mml:mover></mml:math></inline-formula> is the mean of the true temperature values over the evaluation set. MAE provides an interpretable error in degrees Celsius, while RMSE penalizes larger errors more heavily. R<sup>2</sup> indicates how well the predictions track the true temperature trajectory relative to the mean; a value below zero indicates performance worse than simply predicting the mean temperature.</p>
      </sec>
    </sec>
    <sec id="sec4">
      <title>4. Results</title>
      <sec id="sec4dot1">
        <title>4.1. In-Distribution Performance (Dataset A)</title>
        <p>Both models were trained exclusively on Dataset A and evaluated first on the same dataset to establish their baseline fitting capability. The DNN achieved an MAE of 0.066˚C, RMSE of 0.077˚C, and R<sup>2</sup> of 0.9999, indicating near-perfect reconstruction of the training temperature trajectory. The PINN achieved an MAE of 0.079˚C, RMSE of 0.108˚C, and R<sup>2</sup> of 0.9997. Both models fit the training data well, but the DNN fits more tightly, as expected. When a neural network is given sufficient clean data from the same distribution it was trained on, the physics constraint introduces a competing objective that prevents the model from fitting the data as closely as an unconstrained network. This small performance gap on in-distribution data is consistent with the general behavior of regularized versus unregularized models and does not indicate a problem with the PINN formulation.</p>
      </sec>
      <sec id="sec4dot2">
        <title>4.2. Moderate Distribution Shift (Dataset B1)</title>
        <p>Dataset B1 represents a change in load step timing and magnitude that the models were never exposed to during training. Under this condition, the DNN MAE increased from 0.066˚C to 0.434˚C—a degradation of approximately 6.6 times relative to in-distribution performance. The R<sup>2</sup> remained at 0.9939, indicating that while the DNN’s absolute errors increased, it still tracked the general temperature trend reasonably well. The PINN produced an MAE of 0.471˚C and R<sup>2</sup> of 0.9926 under the same condition, slightly worse than the DNN on both metrics.</p>
        <p>A systematic sweep of the physics loss weight <italic>λ</italic> from 1 × 10<sup>−</sup><sup>6</sup> to 1 × 10<sup>−</sup><sup>3</sup> was conducted to determine whether any value of <italic>λ</italic> could produce a PINN advantage over the DNN. The results are presented in <bold>Table 2</bold>. Across all tested values, PINN performance degraded monotonically as <italic>λ</italic> increased on both Dataset A and Dataset B1. No value of <italic>λ</italic> produced a PINN that outperformed the DNN. At <italic>λ</italic> = 1 × 10<sup>−</sup><sup>6</sup>, the physics constraint had negligible effect and the PINN matched the DNN almost exactly. At <italic>λ</italic> = 1 × 10<sup>−</sup><sup>3</sup>, both in-distribution and shift performance degraded substantially. This confirms that the DNN superiority on Dataset B1 is not an artifact of suboptimal hyperparameter selection. This trend is visually summarized in <xref ref-type="fig" rid="fig2">Figure 2</xref>, which plots the MAE against the physics loss weight <italic>λ</italic><italic>.</italic></p>
        <p><bold>Table 2</bold><bold>.</bold> Lambda sensitivity analysis—PINN performance vs. physics loss weight.</p>
        <table-wrap id="tbl2">
          <label>Table 2</label>
          <table>
            <tbody>
              <tr>
                <td>
                  Lambda (
                  <italic>λ</italic>
                  )
                </td>
                <td>A MAE (˚C)</td>
                <td>
                  A (R
                  <sup>2</sup>
                  )
                </td>
                <td>B1 MAE (˚C)</td>
                <td>
                  B1 (R
                  <sup>2</sup>
                  )
                </td>
              </tr>
              <tr>
                <td>DNN baseline</td>
                <td>0.0660</td>
                <td>0.9999</td>
                <td>0.4338</td>
                <td>0.9939</td>
              </tr>
              <tr>
                <td>
                  1 × 10
                  <sup>−</sup>
                  <sup>6</sup>
                </td>
                <td>0.0660</td>
                <td>0.9999</td>
                <td>0.4342</td>
                <td>0.9939</td>
              </tr>
              <tr>
                <td>
                  5 × 10
                  <sup>−</sup>
                  <sup>6</sup>
                </td>
                <td>0.0660</td>
                <td>0.9999</td>
                <td>0.4356</td>
                <td>0.9938</td>
              </tr>
              <tr>
                <td>
                  1 × 10
                  <sup>−</sup>
                  <sup>5</sup>
                </td>
                <td>0.0661</td>
                <td>0.9999</td>
                <td>0.4375</td>
                <td>0.9938</td>
              </tr>
              <tr>
                <td>
                  5 × 10
                  <sup>−</sup>
                  <sup>5</sup>
                </td>
                <td>0.0694</td>
                <td>0.9998</td>
                <td>0.4524</td>
                <td>0.9932</td>
              </tr>
              <tr>
                <td>
                  1 × 10
                  <sup>−</sup>
                  <sup>4</sup>
                </td>
                <td>0.0785</td>
                <td>0.9997</td>
                <td>0.4710</td>
                <td>0.9926</td>
              </tr>
              <tr>
                <td>
                  5 × 10
                  <sup>−</sup>
                  <sup>4</sup>
                </td>
                <td>0.2230</td>
                <td>0.9975</td>
                <td>0.6102</td>
                <td>0.9871</td>
              </tr>
              <tr>
                <td>
                  1 × 10
                  <sup>−</sup>
                  <sup>3</sup>
                </td>
                <td>0.3824</td>
                <td>0.9925</td>
                <td>0.7707</td>
                <td>0.9792</td>
              </tr>
            </tbody>
          </table>
        </table-wrap>
        <fig id="fig2">
          <label>Figure 2</label>
          <graphic xlink:href="https://html.scirp.org/file/2970142-rId42.jpeg?20260825114556" />
        </fig>
        <p><bold>Figure 2</bold><bold>.</bold> PINN MAE as a function of physics loss weight <italic>λ</italic> on Dataset A and Dataset B1, with the DNN baseline shown as dashed lines.</p>
      </sec>
      <sec id="sec4dot3">
        <title>4.3. Severe Distribution Shift (Dataset B2)</title>
        <p>Dataset B2 introduced a more substantial distributional change by elevating the ambient temperature from 25˚C to 40˚C while keeping the load profile identical to Dataset A. Under this condition, the DNN MAE rose to 1.850˚C and R<sup>2</sup> fell to 0.9270, a considerably larger degradation than observed under the moderate load shift of Dataset B1. This confirms that the ambient temperature shift created a more challenging test condition. The PINN produced an MAE of 1.871˚C and R<sup>2</sup> of 0.9247. The gap between DNN and PINN remained small, with the DNN still marginally outperforming the PINN. Notably, R<sup>2</sup> values for both models dropped below 0.93, indicating meaningful degradation in prediction quality under this more severe shift. However, neither model showed a clear advantage, suggesting that under this type of thermal operating point shift, the simplified first-order prior does not provide sufficient constraint to counteract the degradation caused by the unseen ambient condition.</p>
      </sec>
      <sec id="sec4dot4">
        <title>4.4. Cold Start Transient and Noisy Measurements</title>
        <p>Dataset C captured the first 100 seconds of the same loading scenario as Dataset A, representing the early thermal transient from ambient temperature. The DNN achieved an MAE of 0.063˚C and an R<sup>2</sup> of 0.9996, slightly better than its full-duration in-distribution performance. The PINN achieved an MAE of 0.079˚C and an R<sup>2</sup> of 0.9991. The DNN’s superior performance on Dataset C is consistent with the in-distribution finding; the early transient falls within the range of conditions seen during training, and the DNN interpolates accurately without requiring physical guidance. The physics prior does not improve performance when the test condition is well covered by the training data.</p>
        <p>Dataset D introduced 5% Gaussian noise to the speed and current measurement channels, representing sensor uncertainty in a practical deployment scenario. The DNN MAE increased modestly from 0.066˚C to 0.128˚C, while R<sup>2</sup> remained at 0.9994. The PINN produced an MAE of 0.131˚C and R<sup>2</sup> of 0.9994. Both models showed similar robustness to the noise level tested, with neither gaining a measurable advantage. The physics prior did not function as a noise filter at this noise level. This may be because the 5% noise, while detectable, did not substantially alter the underlying input-output mapping that the DNN had learned from clean data.</p>
      </sec>
      <sec id="sec4dot5">
        <title>4.5. Hard Distribution Shift Evaluation</title>
        <p>To test the limits of both models under extreme conditions not represented in the training data, three additional datasets (E1-E3) were evaluated. Dataset E1 elevated the ambient temperature to 50˚C, a 25˚C increase from the training ambient of 25˚C. Dataset E2 applied a 2× load scaling, subjecting the motor to severe mechanical overload. Dataset E3 extended the simulation duration to 500 seconds, capturing cumulative thermal buildup beyond the 200-second training window.</p>
        <p>The results are presented in <bold>Table 3</bold>. Under E1 (50˚C ambient), the DNN achieved an MAE of 13.63˚C while the PINN achieved 13.77˚C. Under E2 (2× load), both models produced MAEs of 13.63˚C and 13.77˚C, respectively. Under E3 (500 s extended run), the DNN achieved an MAE of 27.85˚C and the PINN achieved 28.79˚C.</p>
        <p><bold>Table 3</bold><bold>.</bold> Hard distribution shift results.</p>
        <table-wrap id="tbl3">
          <label>Table 3</label>
          <table>
            <tbody>
              <tr>
                <td>Dataset</td>
                <td>Condition</td>
                <td>DNN MAE (˚C)</td>
                <td>PINN MAE (˚C)</td>
              </tr>
              <tr>
                <td>E1</td>
                <td>50˚C ambient</td>
                <td>13.63</td>
                <td>13.77</td>
              </tr>
              <tr>
                <td>E2</td>
                <td>2× load scaling</td>
                <td>13.63</td>
                <td>13.77</td>
              </tr>
              <tr>
                <td>E3</td>
                <td>500s extended run</td>
                <td>27.85</td>
                <td>28.79</td>
              </tr>
            </tbody>
          </table>
        </table-wrap>
        <p>Both models degraded substantially under these extreme shifts, with MAE increasing significantly relative to in-distribution performance. Critically, the physics prior provided no protective advantage; the DNN and PINN performed equivalently within the margin of measurement uncertainty.</p>
      </sec>
      <sec id="sec4dot6">
        <title>4.6. Sampling Robustness and Data Degradation</title>
        <p>To address the question of whether physics-informed regularization provides value under practical data limitations, four degraded-data scenarios (E4 - E7) were evaluated. Dataset E4 employed sparse sampling at 0.5 s intervals (2 Hz) instead of the original 0.1 s (10 Hz). Dataset E5 simulated missing data by randomly removing 20% of training samples. Dataset E6 introduced biased coverage by training exclusively on high-torque operating regions. Dataset E7 added 10% Gaussian noise to input measurements.</p>
        <p>The results are presented in <bold>Table 4</bold>. The PINN outperformed the DNN across all four degraded-data conditions. Under E4 (sparse sampling), the PINN achieved an MAE of 15.44˚C compared to the DNN MAE of 17.21˚C, a 10.3% improvement. Under E5 (missing data), the PINN MAE was 15.43˚C versus the DNN 17.22˚C (10.4% improvement). Under E6 (biased coverage), the PINN MAE was 11.87˚C versus the DNN 13.33˚C (11.0% improvement). Under E7 (noisy inputs), the PINN MAE was 15.45˚C versus the DNN 17.22˚C (10.3% improvement).</p>
        <p><bold>Table 4</bold><bold>.</bold> Sampling robustness results.</p>
        <table-wrap id="tbl4">
          <label>Table 4</label>
          <table>
            <tbody>
              <tr>
                <td>Condition</td>
                <td>DNN MAE (˚C)</td>
                <td>PINN MAE (˚C)</td>
                <td>Ratio (PINN/DNN)</td>
              </tr>
              <tr>
                <td>E4: Sparse sampling (0.5 s)</td>
                <td>17.21</td>
                <td>15.44</td>
                <td>0.897</td>
              </tr>
              <tr>
                <td>E5: Missing data (20%)</td>
                <td>17.22</td>
                <td>15.43</td>
                <td>0.896</td>
              </tr>
              <tr>
                <td>E6: Biased coverage</td>
                <td>13.33</td>
                <td>11.87</td>
                <td>0.891</td>
              </tr>
              <tr>
                <td>E7: Noisy inputs (10%)</td>
                <td>17.22</td>
                <td>15.45</td>
                <td>0.897</td>
              </tr>
            </tbody>
          </table>
        </table-wrap>
        <p>This finding establishes a critical boundary condition: the physics prior provides measurable benefit (approximately 10% MAE reduction) precisely when data quality degrades. Under conditions of complete, clean, and temporally dense data, the unconstrained DNN is sufficient. However, when data are sparse, missing, biased, or noisy, the physics-informed regularizer compensates for information loss.</p>
      </sec>
      <sec id="sec4dot7">
        <title>4.7. Self-Supervised PINN</title>
        <p>A self-supervised PINN was trained using only the physics residual loss, with no temperature labels provided during training. After 3000 epochs, the physics loss converged to near zero, indicating that the network found a solution that approximately satisfies the thermal ODE as shown in the training loss curves in <xref ref-type="fig" rid="fig3">Figure 3</xref>. However, evaluation against true temperature labels revealed MAE values of approximately 30˚C across all datasets, with R<sup>2</sup> values below −18. A negative R<sup>2</sup> of this magnitude indicates that the model performs far worse than simply predicting the mean temperature; the physically consistent solution found by the optimizer does not correspond to the actual thermal trajectory of the motor.</p>
        <p>This result confirms that the first-order thermal ODE, while thermodynamically valid, is underdetermined as a standalone estimator. The ODE admits multiple solutions that satisfy the residual criterion; the network converges to one of these solutions, but without temperature labels to anchor it, there is no guarantee that this solution matches physical reality. Self-supervised thermal estimation using a simplified first-order prior is therefore not viable with the current model structure.</p>
        <fig id="fig3">
          <label>Figure 3</label>
          <graphic xlink:href="https://html.scirp.org/file/2970142-rId43.jpeg?20260825114557" />
        </fig>
        <p><bold>Figure 3</bold><bold>.</bold> Training loss convergence curves for DNN, PINN, and self-supervised PINN over 3,000 epochs.</p>
      </sec>
      <sec id="sec4dot8">
        <title>4.8. Summary of SPIM Results</title>
        <p><bold>Table 5</bold> presents a complete summary of all SPIM results on datasets A-D and the self-supervised experiment. <xref ref-type="fig" rid="fig4">Figure 4</xref> provides a comparative bar chart of the MAE across these conditions.</p>
        <p><bold>Table 5</bold><bold>.</bold> SPIM Results—DNN vs. PINN across all conditions.</p>
        <table-wrap id="tbl5">
          <label>Table 5</label>
          <table>
            <tbody>
              <tr>
                <td>Dataset</td>
                <td>Condition</td>
                <td>DNN (MAE)</td>
                <td>
                  DNN (R
                  <sup>2</sup>
                  )
                </td>
                <td>PINN (MAE)</td>
                <td>
                  PINN (R
                  <sup>2</sup>
                  )
                </td>
                <td>Winner</td>
              </tr>
              <tr>
                <td>A</td>
                <td>Training</td>
                <td>0.0660</td>
                <td>0.9999</td>
                <td>0.0785</td>
                <td>0.9997</td>
                <td>DNN</td>
              </tr>
              <tr>
                <td>B1</td>
                <td>Moderate load shift</td>
                <td>0.4338</td>
                <td>0.9939</td>
                <td>0.4710</td>
                <td>0.9926</td>
                <td>DNN</td>
              </tr>
              <tr>
                <td>B2</td>
                <td>Severe ambient shift</td>
                <td>1.8500</td>
                <td>0.9270</td>
                <td>1.8705</td>
                <td>0.9247</td>
                <td>DNN</td>
              </tr>
              <tr>
                <td>C</td>
                <td>Cold start</td>
                <td>0.0629</td>
                <td>0.9996</td>
                <td>0.0786</td>
                <td>0.9991</td>
                <td>DNN</td>
              </tr>
              <tr>
                <td>D</td>
                <td>Sensor noise</td>
                <td>0.1276</td>
                <td>0.9994</td>
                <td>0.1313</td>
                <td>0.9994</td>
                <td>DNN</td>
              </tr>
              <tr>
                <td>Self-supervised</td>
                <td>No labels</td>
                <td>—</td>
                <td>—</td>
                <td>~30</td>
                <td>&lt; −18</td>
                <td>Fails</td>
              </tr>
            </tbody>
          </table>
        </table-wrap>
        <fig id="fig4">
          <label>Figure 4</label>
          <graphic xlink:href="https://html.scirp.org/file/2970142-rId44.jpeg?20260825114557" />
        </fig>
        <p><bold>Figure 4</bold><bold>.</bold> MAE comparison between DNN and PINN across all evaluation datasets (<italic>λ</italic> = 1 × 10<sup>−</sup><sup>4</sup>).</p>
        <p>The evaluation on datasets E1 - E7, summarized in <bold>Table 4</bold> and <bold>Table 5</bold>, further clarifies the boundary conditions of physics prior effectiveness. Under extreme distribution shifts (E1 - E3), both models degraded substantially and equally, with MAE increasing to 13.6˚C - 28.8˚C and the physics prior providing no protective advantage. However, under degraded data scenarios (E4 - E7), the PINN consistently outperformed the DNN, achieving approximately 10% lower MAE. This establishes a critical boundary condition: the physics prior becomes valuable when data quality degrades, while offering no benefit under complete data conditions or extreme distribution shifts.</p>
      </sec>
    </sec>
    <sec id="sec5">
      <title>5. Discussion</title>
      <sec id="sec5dot1">
        <title>5.1. Boundary Conditions for Physics-Prior Effectiveness</title>
        <p>The results, summarized in <bold>Tables 3</bold><bold>-</bold><bold>5</bold>, reveal a nuanced picture of physics-informed neural network effectiveness. The findings can be categorized into two distinct regimes:</p>
        <p><bold>Regime 1</bold><bold>—</bold><bold>Complete Data with Temporal Context</bold>: When temperature labels are abundant, sampling is dense, data coverage is unbiased, and time is available as an input feature, DNN and PINN achieve equivalent accuracy (datasets A–D). The physics prior provides no measurable advantage because the data alone contain sufficient information to learn the thermal mapping.</p>
        <p><bold>Regime 2</bold><bold>—</bold><bold>Degraded Data</bold>: When data quality degrades (sparse sampling, missing observations, biased operating coverage, or noisy measurements—datasets E4 - E7), the physics prior provides meaningful regularization, reducing MAE by approximately 10% relative to the DNN alone. This finding demonstrates that physics-informed learning is valuable not as a universal improvement but as a targeted tool for data-limited scenarios.</p>
        <p>Under extreme distribution shifts (E1 - E3), both models degraded substantially and equally. The physics prior offered no protection against extrapolation failure, consistent with the explanation that the first-order ODE, while encoding thermodynamic relationships, does not capture all relevant physical mechanisms (e.g., core iron losses, rotor copper losses, multi-dimensional heat paths).</p>
      </sec>
      <sec id="sec5dot2">
        <title>5.2. Why the Physics Prior Did Not Help</title>
        <p>There are three reasons why the simplified thermal ODE failed to provide a measurable benefit under any tested conditions.</p>
        <p>First, the distribution shifts tested here, while genuine, did not degrade the DNN’s interpolation capability to the point of failure. On Dataset B1, the DNN achieved an R<sup>2</sup> of 0.9939 even under the unseen load profile. On Dataset B2, R<sup>2</sup> fell to 0.9270, a meaningful degradation, but still far from the regime where a data-driven model breaks down entirely. The physics prior is most useful when a model is forced to extrapolate far outside its training distribution. In this study, both shift conditions were moderate enough that the DNN retained reasonable predictive capability without physical guidance.</p>
        <p>Second, the simplified first-order ODE is an incomplete representation of the motor’s actual thermal dynamics. The thermal prior used here accounts for copper losses in both windings, convective cooling, and bearing friction losses following Calasan <italic>et al</italic>. [<xref ref-type="bibr" rid="B14">14</xref>]. However, it excludes core iron losses, rotor copper losses, thermal interface resistances, and multi-dimensional heat flow paths. When the physics prior does not capture the dominant thermal mechanisms accurately, it does not guide the model toward better predictions; it guides it toward a thermodynamically approximate but empirically inaccurate solution. This is precisely what the lambda sensitivity analysis revealed: as the physics weight increased, both in-distribution and shift performance degraded monotonically. There was no operating point where physics helped.</p>
        <p>Third, the self-supervised experiment confirmed that the ODE prior is underdetermined. A network trained purely on the physics residual converged to a near-zero residual loss but produced temperature predictions with an MAE of approximately 30˚C. This demonstrates that multiple solutions satisfy the simplified ODE; the physics constraint alone cannot uniquely identify the motor’s thermal trajectory. Temperature labels are necessary to anchor the solution, and when they are present in sufficient quantity, the data-driven model uses them effectively without needing physical guidance.</p>
      </sec>
      <sec id="sec5dot3">
        <title>5.3. Comparison with Prior Work</title>
        <p>The finding here contrasts with results reported by Wang <italic>et al</italic>. [<xref ref-type="bibr" rid="B10">10</xref>], who demonstrated a physics-informed approach achieving an MAE of 1.71˚C on real PMSM data. However, there is an important difference in approach. Wang <italic>et al</italic>. used a 10-node Motor-CAD-calibrated lumped-parameter thermal network, a high-fidelity multi-node prior that captures inter-component heat transfer, not a single first-order ODE. A more faithful physics prior naturally provides stronger guidance and is more likely to yield measurable benefits. The present study tests the other end of this spectrum, a minimal first-order prior, and the results suggest that the benefit of physics-informed regularization is strongly dependent on the fidelity of the physics model used.</p>
        <p>Stensson [<xref ref-type="bibr" rid="B11">11</xref>] and Lim <italic>et al</italic>. [<xref ref-type="bibr" rid="B12">12</xref>] similarly demonstrate PINN benefits under specific conditions, but neither study isolates the marginal contribution of the physics prior by comparing against an identical unconstrained baseline under the same conditions. The controlled experimental design in the present work, with identical architecture and data and only the loss function differing, allows a cleaner attribution of performance differences specifically to the physics constraint.</p>
      </sec>
      <sec id="sec5dot4">
        <title>5.4. The Single-Phase Induction Motor—An Underexplored Domain</title>
        <p>Prior PINN work on motor thermal estimation has focused almost exclusively on PMSMs [<xref ref-type="bibr" rid="B10">10</xref>][<xref ref-type="bibr" rid="B11">11</xref>][<xref ref-type="bibr" rid="B16">16</xref>]. The single-phase induction motor, despite its widespread use in household appliances, small pumps, and light industrial equipment, has received almost no attention in the neural network-based thermal estimation literature. The dual-winding structure of the SPIM, with separate main and auxiliary windings, each contributing to copper losses, requires a more nuanced thermal prior than the single-winding models commonly used for PMSMs. This study provides the first controlled diagnostic comparison of DNN and PINN for SPIM thermal estimation and establishes the thermal ODE formulation for this motor type with dual-winding copper loss terms.</p>
      </sec>
      <sec id="sec5dot5">
        <title>5.5. Practical Implications</title>
        <p>The results of this study convey a clear practical message for engineers deploying motor thermal estimators.</p>
        <p>When training data are available and cover the intended operating range, a well-trained feedforward DNN is the preferred choice. It is simpler to implement, easier to train, requires no motor-specific thermal parameters, and consistently achieves lower prediction error than a PINN using a simplified physics prior.</p>
        <p>The physics prior adds value only when the simplified ODE is a faithful approximation of the dominant thermal dynamics and when the distribution shift is severe enough to substantially degrade data-driven interpolation. Neither condition was met in this study. For applications where these conditions are expected, such as motors operating under highly variable or unpredictable duty cycles with limited training data, a higher-fidelity prior, such as a multi-node LPTN calibrated to the specific motor, would be more appropriate than the first-order ODE used here.</p>
        <p>Self-supervised thermal estimation, training without temperature labels using only the physics loss, is not viable with a simplified first-order prior. The underdetermined nature of the ODE means that the network converges to physically plausible but empirically incorrect solutions. Applications that require label-free deployment will need either higher-fidelity physics models or alternative approaches, such as transfer learning or domain adaptation.</p>
      </sec>
      <sec id="sec5dot6">
        <title>5.6. Limitations</title>
        <p>Several limitations of this study should be noted.</p>
        <p>First, the SPIM datasets were generated from a Simulink simulation model parameterized from manufacturer datasheet values. While the model produces physically plausible thermal trajectories, confirmed by steady-state temperature, speed, and torque values within expected ranges for a 1 HP motor, it has not been validated against direct thermocouple measurements from a physical motor. The absolute MAE values reported here reflect simulation-to-simulation performance and should not be compared directly to results from studies using real hardware measurements. Validation on real motor hardware remains an important direction for future work.</p>
        <p>Second, the distribution shifts tested, while covering a range of operating conditions, were moderate in severity. The most severe shift tested, an ambient temperature change from 25˚C to 40˚C, degraded DNN R<sup>2</sup> to 0.927, which still represents reasonable predictive capability. Under more extreme shifts, such as a change in motor drive topology, significantly different ambient environments, or operation in completely different load regimes, the physics prior might provide a measurable advantage. This represents a direction for future investigation.</p>
        <p>Third, the network architecture was fixed throughout the study (761 parameters). Architecture sensitivity was not explored; deeper or wider networks may alter the observed boundary conditions. Future work should examine whether network capacity affects the relative performance of DNN and PINN.</p>
        <p>Fourth, the lambda sensitivity analysis fixed all other hyperparameters (optimizer, learning rate, loss scaling). Alternative optimizer settings and loss-balancing strategies such as GradNorm [<xref ref-type="bibr" rid="B18">18</xref>] merit future investigation.</p>
        <p>Fifth, the study focused on a single motor type (SPIM). Generalization of these findings to three-phase induction motors, brushless DC motors, or switched reluctance motors requires additional studies.</p>
      </sec>
    </sec>
    <sec id="sec6">
      <title>6. Conclusions</title>
      <p>This controlled diagnostic study demonstrates that the effectiveness of a first-order thermal physics prior for motor thermal virtual sensing is conditional, not universal. When temperature labels are abundant and temporal context is available, physics-informed and data-driven networks achieve equivalent accuracy. When data quality degrades, the physics prior provides meaningful regularization.</p>
      <p>These findings establish a practical decision boundary for engineers deploying motor thermal estimators. When training data adequately covers the intended operating range, a well-trained DNN is the preferred choice. When data quality is compromised by sparse sampling, missing observations, biased coverage, or measurement noise, a physics-informed approach offers modest but meaningful improvements.</p>
      <p>Future work should validate these findings on real motor hardware with direct thermocouple measurements, explore architecture sensitivity across network sizes and depths, and investigate whether higher-fidelity multi-node thermal priors provide greater benefits under extreme shifts.</p>
    </sec>
    <sec id="sec7">
      <title>Author Contributions</title>
      <p>Conceptualization, T.A.A. and A.D.C.; methodology, T.A.A. and S.Z.Z.; software, T.A.A. and D.O.O.; validation, T.A.A., M.A.M., and E.S.O. formal analysis, T.A.A. and S.Z.Z.; investigation, D.O.O. and E.S.O.; resources, A.D.C. and S.Z.Z.; data curation, T.A.A. and D.O.O.; writing—original draft preparation, T.A.A.; writing—review and editing, A.D.C., S.Z.Z., and M.A.M.; visualization, T.A.A. and E.S.O.; supervision, A.D.C. and S.Z.Z.; project administration, A.D.C. All authors have read and agreed to the published version of the manuscript.</p>
    </sec>
  </body>
  <back>
    <ref-list>
      <title>References</title>
      <ref id="B1">
        <label>1.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Mostajeran, E., Amiri, N., Ebrahimi, S. and Jatskevich, J. (2023) Electrical Machines in Electromagnetic Transient Simulations: Focusing on Efficient and Accurate Models. <italic>IEEE</italic><italic>Electrification</italic><italic>Magazine</italic>, 11, 38-53. https://doi.org/10.1109/mele.2023.3320508 <pub-id pub-id-type="doi">10.1109/mele.2023.3320508</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/mele.2023.3320508">https://doi.org/10.1109/mele.2023.3320508</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Mostajeran, E.</string-name>
              <string-name>Amiri, N.</string-name>
              <string-name>Ebrahimi, S.</string-name>
              <string-name>Jatskevich, J.</string-name>
            </person-group>
            <year>2023</year>
            <article-title>Electrical Machines in Electromagnetic Transient Simulations: Focusing on Efficient and Accurate Models</article-title>
            <source>IEEE Electrification Magazine</source>
            <volume>11</volume>
            <pub-id pub-id-type="doi">10.1109/mele.2023.3320508</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B2">
        <label>2.</label>
        <citation-alternatives>
          <mixed-citation publication-type="confproc">Jing, H., Xiao, D., Wang, X., Chen, Z., Fang, G. and Guo, X. (2022) Temperature Estimation of Permanent Magnet Synchronous Motors Using Support Vector Regression. 2022 25 <italic>th International Conference on Electrical Machines and Systems</italic>( <italic>ICEMS</italic>), Chiang Mai, 29 November-2 December 2022, 1-6. https://doi.org/10.1109/icems56177.2022.9983067 <pub-id pub-id-type="doi">10.1109/icems56177.2022.9983067</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/icems56177.2022.9983067">https://doi.org/10.1109/icems56177.2022.9983067</ext-link></mixed-citation>
          <element-citation publication-type="confproc">
            <person-group person-group-type="author">
              <string-name>Jing, H.</string-name>
              <string-name>Xiao, D.</string-name>
              <string-name>Wang, X.</string-name>
              <string-name>Chen, Z.</string-name>
              <string-name>Fang, G.</string-name>
              <string-name>Guo, X.</string-name>
            </person-group>
            <year>2022</year>
            <article-title>Temperature Estimation of Permanent Magnet Synchronous Motors Using Support Vector Regression</article-title>
            <source>2022 25th International Conference on Electrical Machines and Systems (ICEMS)</source>
            <volume>29</volume>
            <pub-id pub-id-type="doi">10.1109/icems56177.2022.9983067</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B3">
        <label>3.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Boldea, I. (2017) Electric Generators and Motors: An Overview. <italic>CES Transactions on Electrical Machines and Systems</italic>, 1, 3-14. https://doi.org/10.23919/tems.2017.7911104 <pub-id pub-id-type="doi">10.23919/tems.2017.7911104</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.23919/tems.2017.7911104">https://doi.org/10.23919/tems.2017.7911104</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Boldea, I.</string-name>
            </person-group>
            <year>2017</year>
            <article-title>Electric Generators and Motors: An Overview</article-title>
            <source>CES Transactions on Electrical Machines and Systems</source>
            <volume>1</volume>
            <pub-id pub-id-type="doi">10.23919/tems.2017.7911104</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B4">
        <label>4.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Kim, D.W., Kang, D.H., Kim, C.H., Kim, J.S., Kim, Y.J. and Jung, S.Y. (2020) Operation Characteristic of IPMSM Considering PM Saturation Temperature. <italic>IEEE</italic><italic>Trans</italic><italic>actions</italic><italic>on</italic><italic>Applied</italic><italic>Superconductivity</italic>, 30, 1-4. https://doi.org/10.1109/tasc.2020.2989799 <pub-id pub-id-type="doi">10.1109/tasc.2020.2989799</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/tasc.2020.2989799">https://doi.org/10.1109/tasc.2020.2989799</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Kim, D.W.</string-name>
              <string-name>Kang, D.H.</string-name>
              <string-name>Kim, C.H.</string-name>
              <string-name>Kim, J.S.</string-name>
              <string-name>Kim, Y.J.</string-name>
              <string-name>Jung, S.Y.</string-name>
            </person-group>
            <year>2020</year>
            <article-title>Operation Characteristic of IPMSM Considering PM Saturation Temperature</article-title>
            <source>IEEE Transactions on Applied Superconductivity</source>
            <volume>30</volume>
            <pub-id pub-id-type="doi">10.1109/tasc.2020.2989799</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B5">
        <label>5.</label>
        <citation-alternatives>
          <mixed-citation publication-type="journal">Wallscheid, O. (2021) Thermal Monitoring of Electric Motors: State-of-the-Art Review and Future Challenges. <italic>IEEE</italic><italic>Open</italic><italic>Journal</italic><italic>of</italic><italic>Industry</italic><italic>Applications</italic>, 2, 204-223. https://doi.org/10.1109/ojia.2021.3091870 <pub-id pub-id-type="doi">10.1109/ojia.2021.3091870</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/ojia.2021.3091870">https://doi.org/10.1109/ojia.2021.3091870</ext-link></mixed-citation>
          <element-citation publication-type="journal">
            <person-group person-group-type="author">
              <string-name>Wallscheid, O.</string-name>
            </person-group>
            <year>2021</year>
            <article-title>Thermal Monitoring of Electric Motors: State-of-the-Art Review and Future Challenges</article-title>
            <source>IEEE Open Journal of Industry Applications</source>
            <volume>2</volume>
            <pub-id pub-id-type="doi">10.1109/ojia.2021.3091870</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B6">
        <label>6.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Wallscheid, O. and Bocker, J. (2016) Global Identification of a Low-Order Lumped-Parameter Thermal Network for Permanent Magnet Synchronous Motors. <italic>IEEE</italic><italic>Transactions</italic><italic>on</italic><italic>Energy</italic><italic>Conversion</italic>, 31, 354-365. https://doi.org/10.1109/tec.2015.2473673 <pub-id pub-id-type="doi">10.1109/tec.2015.2473673</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/tec.2015.2473673">https://doi.org/10.1109/tec.2015.2473673</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Wallscheid, O.</string-name>
              <string-name>Bocker, J.</string-name>
            </person-group>
            <year>2016</year>
            <article-title>Global Identification of a Low-Order Lumped-Parameter Thermal Network for Permanent Magnet Synchronous Motors</article-title>
            <source>IEEE Transactions on Energy Conversion</source>
            <volume>31</volume>
            <pub-id pub-id-type="doi">10.1109/tec.2015.2473673</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B7">
        <label>7.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Kral, C., Haumer, A. and Lee, S.B. (2014) A Practical Thermal Model for the Estimation of Permanent Magnet and Stator Winding Temperatures. <italic>IEEE</italic><italic>Transactions</italic><italic>on</italic><italic>Power</italic><italic>Electronics</italic>, 29, 455-464. https://doi.org/10.1109/tpel.2013.2253128 <pub-id pub-id-type="doi">10.1109/tpel.2013.2253128</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/tpel.2013.2253128">https://doi.org/10.1109/tpel.2013.2253128</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Kral, C.</string-name>
              <string-name>Haumer, A.</string-name>
              <string-name>Lee, S.B.</string-name>
            </person-group>
            <year>2014</year>
            <article-title>A Practical Thermal Model for the Estimation of Permanent Magnet and Stator Winding Temperatures</article-title>
            <source>IEEE Transactions on Power Electronics</source>
            <volume>29</volume>
            <pub-id pub-id-type="doi">10.1109/tpel.2013.2253128</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B8">
        <label>8.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Jin, L., Mao, Y., Wang, X., Lu, L. and Wang, Z. (2024) A Model-Based and Data-Driven Integrated Temperature Estimation Method for PMSM. <italic>IEEE</italic><italic>Transactions</italic><italic>on</italic><italic>Power</italic><italic>Electronics</italic>, 39, 8553-8561. https://doi.org/10.1109/tpel.2024.3382300 <pub-id pub-id-type="doi">10.1109/tpel.2024.3382300</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/tpel.2024.3382300">https://doi.org/10.1109/tpel.2024.3382300</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Jin, L.</string-name>
              <string-name>Mao, Y.</string-name>
              <string-name>Wang, X.</string-name>
              <string-name>Lu, L.</string-name>
              <string-name>Wang, Z.</string-name>
            </person-group>
            <year>2024</year>
            <article-title>A Model-Based and Data-Driven Integrated Temperature Estimation Method for PMSM</article-title>
            <source>IEEE Transactions on Power Electronics</source>
            <volume>39</volume>
            <pub-id pub-id-type="doi">10.1109/tpel.2024.3382300</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B9">
        <label>9.</label>
        <citation-alternatives>
          <mixed-citation publication-type="journal">Raissi, M., Perdikaris, P. and Karniadakis, G.E. (2019) Physics-Informed Neural Networks: A Deep Learning Framework for Solving Forward and Inverse Problems Involving Nonlinear Partial Differential Equations. <italic>Journal</italic><italic>of</italic><italic>Computational</italic><italic>Physics</italic>, 378, 686-707. https://doi.org/10.1016/j.jcp.2018.10.045 <pub-id pub-id-type="doi">10.1016/j.jcp.2018.10.045</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1016/j.jcp.2018.10.045">https://doi.org/10.1016/j.jcp.2018.10.045</ext-link></mixed-citation>
          <element-citation publication-type="journal">
            <person-group person-group-type="author">
              <string-name>Raissi, M.</string-name>
              <string-name>Perdikaris, P.</string-name>
              <string-name>Karniadakis, G.E.</string-name>
            </person-group>
            <year>2019</year>
            <article-title>Physics-Informed Neural Networks: A Deep Learning Framework for Solving Forward and Inverse Problems Involving Nonlinear Partial Differential Equations</article-title>
            <source>Journal of Computational Physics</source>
            <volume>378</volume>
            <pub-id pub-id-type="doi">10.1016/j.jcp.2018.10.045</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B10">
        <label>10.</label>
        <citation-alternatives>
          <mixed-citation publication-type="confproc">Wang, P., Wang, X. and Wang, Y. (2023) Physics-Informed Machine Learning Based Permanent Magnet Synchronous Motor Temperature Estimation. 2023 9 <italic>th International Conference on Mechanical and Electronics Engineering</italic>( <italic>ICMEE</italic>), Xi’an, 17-19 November 2023, 582-587. https://doi.org/10.1109/icmee59781.2023.10525605 <pub-id pub-id-type="doi">10.1109/icmee59781.2023.10525605</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/icmee59781.2023.10525605">https://doi.org/10.1109/icmee59781.2023.10525605</ext-link></mixed-citation>
          <element-citation publication-type="confproc">
            <person-group person-group-type="author">
              <string-name>Wang, P.</string-name>
              <string-name>Wang, X.</string-name>
              <string-name>Wang, Y.</string-name>
            </person-group>
            <year>2023</year>
            <article-title>Physics-Informed Machine Learning Based Permanent Magnet Synchronous Motor Temperature Estimation</article-title>
            <source>2023 9th International Conference on Mechanical and Electronics Engineering (ICMEE)</source>
            <volume>17</volume>
            <pub-id pub-id-type="doi">10.1109/icmee59781.2023.10525605</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B11">
        <label>11.</label>
        <citation-alternatives>
          <mixed-citation publication-type="thesis">Stensson, J. and Svantesson, K. (2023) Physics Informed Neural Network for Thermal Modeling of an Electric Motor. Master’s Thesis, Chalmers University of Technology.</mixed-citation>
          <element-citation publication-type="thesis">
            <person-group person-group-type="author">
              <string-name>Stensson, J.</string-name>
              <string-name>Svantesson, K.</string-name>
              <string-name>Thesis, C</string-name>
            </person-group>
            <year>2023</year>
            <article-title>Physics Informed Neural Network for Thermal Modeling of an Electric Motor</article-title>
            <source>Master’s Thesis</source>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B12">
        <label>12.</label>
        <citation-alternatives>
          <mixed-citation publication-type="journal">Lim, H., Lee, J.W., Boyack, J. and Choi, J.B. (2024) EV-PINN: A Physics-Informed Neural Network for Predicting Electric Vehicle Dynamics. arXiv: 2411.14691.</mixed-citation>
          <element-citation publication-type="journal">
            <person-group person-group-type="author">
              <string-name>Lim, H.</string-name>
              <string-name>Lee, J.W.</string-name>
              <string-name>Boyack, J.</string-name>
              <string-name>Choi, J.B.</string-name>
            </person-group>
            <year>2024</year>
            <article-title>EV-PINN: A Physics-Informed Neural Network for Predicting Electric Vehicle Dynamics</article-title>
            <fpage>2411</fpage>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B13">
        <label>13.</label>
        <citation-alternatives>
          <mixed-citation publication-type="journal">Raissi, M., Perdikaris, P. and Karniadakis, G.E. (2017) Physics Informed Deep Learning (Part I): Data-Driven Solutions of Nonlinear Partial Differential Equations. arXiv: 1711.10561.</mixed-citation>
          <element-citation publication-type="journal">
            <person-group person-group-type="author">
              <string-name>Raissi, M.</string-name>
              <string-name>Perdikaris, P.</string-name>
              <string-name>Karniadakis, G.E.</string-name>
            </person-group>
            <year>2017</year>
            <article-title>Physics Informed Deep Learning (Part I): Data-Driven Solutions of Nonlinear Partial Differential Equations</article-title>
            <fpage>1711</fpage>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B14">
        <label>14.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Ćalasan, M., Alqarni, M., Rosić, M., Koljčević, N., Alamri, B. and Abdel Aleem, S.H.E. (2021) A Novel Exact Analytical Solution Based on Kloss Equation towards Accurate Speed-Time Characteristics Modeling of Induction Machines during No-Load Direct Startups. <italic>Applied</italic><italic>Sciences</italic>, 11, Article 5102. https://doi.org/10.3390/app11115102 <pub-id pub-id-type="doi">10.3390/app11115102</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.3390/app11115102">https://doi.org/10.3390/app11115102</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Alqarni, M.</string-name>
              <string-name>Alamri, B.</string-name>
              <string-name>Aleem, S.H.E.</string-name>
            </person-group>
            <year>2021</year>
            <article-title>A Novel Exact Analytical Solution Based on Kloss Equation towards Accurate Speed-Time Characteristics Modeling of Induction Machines during No-Load Direct Startups</article-title>
            <source>Applied Sciences</source>
            <volume>11</volume>
            <elocation-id>5102</elocation-id>
            <pub-id pub-id-type="doi">10.3390/app11115102</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B15">
        <label>15.</label>
        <citation-alternatives>
          <mixed-citation publication-type="confproc">Chowdhury, S.K. and Baski, P.K. (2010) A Simple Lumped Parameter Thermal Model for Electrical Machine of TEFC Design. 2010 <italic>Joint International Conference on Power Electronics</italic>, <italic>Drives and Energy Systems &amp;</italic>2010 <italic>Power India</italic>, New Delhi, 20-23 December 2010, 1-7. https://doi.org/10.1109/pedes.2010.5712385 <pub-id pub-id-type="doi">10.1109/pedes.2010.5712385</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/pedes.2010.5712385">https://doi.org/10.1109/pedes.2010.5712385</ext-link></mixed-citation>
          <element-citation publication-type="confproc">
            <person-group person-group-type="author">
              <string-name>Chowdhury, S.K.</string-name>
              <string-name>Baski, P.K.</string-name>
              <string-name>Electronics, D</string-name>
              <string-name>India, N</string-name>
            </person-group>
            <year>2010</year>
            <article-title>A Simple Lumped Parameter Thermal Model for Electrical Machine of TEFC Design</article-title>
            <source>2010 Joint International Conference on Power Electronics</source>
            <volume>20</volume>
            <pub-id pub-id-type="doi">10.1109/pedes.2010.5712385</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B16">
        <label>16.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Kirchgassner, W., Wallscheid, O. and Bocker, J. (2021) Data-Driven Permanent Magnet Temperature Estimation in Synchronous Motors with Supervised Machine Learning: A Benchmark. <italic>IEEE</italic><italic>Transactions</italic><italic>on</italic><italic>Energy</italic><italic>Conversion</italic>, 36, 2059-2067. https://doi.org/10.1109/tec.2021.3052546 <pub-id pub-id-type="doi">10.1109/tec.2021.3052546</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1109/tec.2021.3052546">https://doi.org/10.1109/tec.2021.3052546</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Kirchgassner, W.</string-name>
              <string-name>Wallscheid, O.</string-name>
              <string-name>Bocker, J.</string-name>
            </person-group>
            <year>2021</year>
            <article-title>Data-Driven Permanent Magnet Temperature Estimation in Synchronous Motors with Supervised Machine Learning: A Benchmark</article-title>
            <source>IEEE Transactions on Energy Conversion</source>
            <volume>36</volume>
            <pub-id pub-id-type="doi">10.1109/tec.2021.3052546</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B17">
        <label>17.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Reinbold, P.A.K., Kageorge, L.M., Schatz, M.F. and Grigoriev, R.O. (2021) Robust Learning from Noisy, Incomplete, High-Dimensional Experimental Data via Physically Constrained Symbolic Regression. <italic>Nature Communications</italic>, 12, Article No. 3219. https://doi.org/10.1038/s41467-021-23479-0 <pub-id pub-id-type="doi">10.1038/s41467-021-23479-0</pub-id><pub-id pub-id-type="pmid">34050155</pub-id><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1038/s41467-021-23479-0">https://doi.org/10.1038/s41467-021-23479-0</ext-link></mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Reinbold, P.A.K.</string-name>
              <string-name>Kageorge, L.M.</string-name>
              <string-name>Schatz, M.F.</string-name>
              <string-name>Grigoriev, R.O.</string-name>
              <string-name>Noisy, I</string-name>
            </person-group>
            <year>2021</year>
            <article-title>Robust Learning from Noisy, Incomplete, High-Dimensional Experimental Data via Physically Constrained Symbolic Regression</article-title>
            <source>Nature Communications</source>
            <volume>12</volume>
            <elocation-id>No</elocation-id>
            <pub-id pub-id-type="doi">10.1038/s41467-021-23479-0</pub-id>
            <pub-id pub-id-type="pmid">34050155</pub-id>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B18">
        <label>18.</label>
        <citation-alternatives>
          <mixed-citation publication-type="journal">Chen, Z., Badrinarayanan, V., Lee, C.Y. and Rabinovich, A. (2017) GradNorm: Gradient Normalization for Adaptive Loss Balancing in Deep Multitask Networks. arXiv: 1711.02257.</mixed-citation>
          <element-citation publication-type="journal">
            <person-group person-group-type="author">
              <string-name>Chen, Z.</string-name>
              <string-name>Badrinarayanan, V.</string-name>
              <string-name>Lee, C.Y.</string-name>
              <string-name>Rabinovich, A.</string-name>
            </person-group>
            <year>2017</year>
            <article-title>GradNorm: Gradient Normalization for Adaptive Loss Balancing in Deep Multitask Networks</article-title>
            <fpage>1711</fpage>
          </element-citation>
        </citation-alternatives>
      </ref>
      <ref id="B19">
        <label>19.</label>
        <citation-alternatives>
          <mixed-citation publication-type="other">Krishnapriyan, A., Gholami, A., Zhe, S., Kirby, R. and Mahoney, M. (2021) Characterizing Possible Failure Modes in Physics-Informed Neural Networks. <italic>Advances in Neural Information Processing Systems</italic>, 34, 26548-26560.</mixed-citation>
          <element-citation publication-type="other">
            <person-group person-group-type="author">
              <string-name>Krishnapriyan, A.</string-name>
              <string-name>Gholami, A.</string-name>
              <string-name>Zhe, S.</string-name>
              <string-name>Kirby, R.</string-name>
              <string-name>Mahoney, M.</string-name>
            </person-group>
            <year>2021</year>
            <article-title>Characterizing Possible Failure Modes in Physics-Informed Neural Networks</article-title>
            <source>Advances in Neural Information Processing Systems</source>
            <volume>34</volume>
          </element-citation>
        </citation-alternatives>
      </ref>
    </ref-list>
  </back>
</article>