\relax 
\emailauthor{cekuo@nchu.edu.tw}{Chih-En Kuo\corref {cor1}}
\Newlabel{cor1}{1}
\Newlabel{label1}{a}
\Newlabel{label2}{b}
\Newlabel{label3}{c}
\citation{Luo2017}
\citation{Ding2025}
\citation{Deng2022}
\citation{Paranjape2023,Zhu2023}
\citation{Ming2023}
\citation{Li2024}
\citation{Sadeghian2025}
\citation{Jiang2024}
\@writefile{toc}{\contentsline {section}{\numberline {1}Introduction}{2}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {1.1}Research Background and Motivation}{2}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {1.2}Research Gap and Problem Statement}{2}{}\protected@file@percent }
\citation{Jin2011}
\citation{AGMA2015}
\citation{Wang2016}
\citation{He2016}
\citation{Hester2018}
\@writefile{toc}{\contentsline {subsection}{\numberline {1.3}Contributions}{3}{}\protected@file@percent }
\citation{ElYousfi2021}
\citation{Luo2017}
\citation{Rao2020,Wu2023}
\@writefile{toc}{\contentsline {section}{\numberline {2}Related Work}{4}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {2.1}Physical Modeling and Analysis of Gear Errors}{4}{}\protected@file@percent }
\citation{Xu2018}
\citation{Chhor2021}
\citation{Feng2020}
\citation{Tsang2025}
\citation{Ming2023}
\citation{Li2024}
\citation{Sadeghian2025}
\citation{Jin2011}
\citation{wrede2024}
\citation{Feng2020}
\@writefile{toc}{\contentsline {subsection}{\numberline {2.2}Data-Driven Optimization Techniques}{5}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {2.3}Reinforcement Learning in Assembly Tasks}{5}{}\protected@file@percent }
\citation{AGMA2015}
\citation{Wang2016}
\citation{AGMA2015}
\@writefile{toc}{\contentsline {section}{\numberline {3}Materials and Methods}{6}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {3.1}System Overview}{6}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {3.2}Data Collection and Feature Engineering}{6}{}\protected@file@percent }
\@writefile{lof}{\contentsline {figure}{\numberline {1}{\ignorespaces Flowchart of the surrogate-assisted system architecture. The framework integrates real-world data boundaries, a Random Forest-based virtual environment simulator, and a Residual Dueling DQN agent, enabling an offline optimization and industrial decision guidance workflow.}}{7}{}\protected@file@percent }
\providecommand*\caption@xref[2]{\@setref\relax\@undefined{#1}}
\newlabel{fig:system_arch}{{1}{7}}
\@writefile{lof}{\contentsline {figure}{\numberline {2}{\ignorespaces Schematic of gear pitch deviations based on ANSI/AGMA standards. The figure illustrates the definitions of single pitch deviation ($f_k$) and cumulative pitch deviation ($f_l$ or $F_p$), which constitute part of the gear-deviation descriptors used in the reinforcement learning state representation.}}{8}{}\protected@file@percent }
\newlabel{fig:gear_pitch}{{2}{8}}
\@writefile{toc}{\contentsline {subsection}{\numberline {3.3}Analysis of Data Characteristics and Optimization Challenges}{8}{}\protected@file@percent }
\citation{He2016}
\citation{espeholt2018impala}
\citation{machado2018revisiting}
\@writefile{toc}{\contentsline {subsection}{\numberline {3.4}Construction of the Data-Driven Virtual Environment}{9}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {3.5}Residual Dueling DQN Architecture}{9}{}\protected@file@percent }
\newlabel{sec:network_arch}{{3.5}{9}}
\@writefile{lof}{\contentsline {figure}{\numberline {3}{\ignorespaces Scatter plots illustrating empirical correlations between physical shim adjustments and the core eight transmission error metrics. The complex, non-linear distribution of data points across both (a) Gear A and (b) Gear B underscores the non-convex, multi-modal nature of the geometric landscape, theoretically justifying a policy-based Deep Reinforcement Learning approach.}}{10}{}\protected@file@percent }
\newlabel{fig:data_distribution}{{3}{10}}
\citation{Wang2016}
\@writefile{lof}{\contentsline {figure}{\numberline {4}{\ignorespaces Architecture of the proposed Residual Dueling DQN for center-distance adjustment. The input state is defined as $s_t=[p_t,z_t,b_t]\in \mathbb  {R}^{13}$, and the output is a Q-value vector $Q(s_t,\cdot )\in \mathbb  {R}^{41}$ corresponding to 41 discrete center-distance adjustment actions.}}{11}{}\protected@file@percent }
\newlabel{fig:dueling_resnet}{{4}{11}}
\@writefile{toc}{\contentsline {subsection}{\numberline {3.6}Proposed Enhancement Modules and Training Strategy}{12}{}\protected@file@percent }
\citation{Schaul2015}
\@writefile{toc}{\contentsline {subsection}{\numberline {3.7}Implementation Details and Hyperparameters}{13}{}\protected@file@percent }
\@writefile{lot}{\contentsline {table}{\numberline {1}{\ignorespaces Hyperparameter configurations for the Residual Dueling DQN framework.}}{14}{}\protected@file@percent }
\newlabel{tab:hyperparameters}{{1}{14}}
\@writefile{toc}{\contentsline {subsection}{\numberline {3.8}Evaluation Metrics}{14}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsubsection}{\numberline {3.8.1}Prediction Accuracy and Error Analysis}{14}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsubsection}{\numberline {3.8.2}Success Definition and Search Efficiency Metrics}{15}{}\protected@file@percent }
\@writefile{toc}{\contentsline {section}{\numberline {4}Experimental Results}{16}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {4.1}Comparative Analysis of Surrogate Models}{16}{}\protected@file@percent }
\@writefile{lot}{\contentsline {table}{\numberline {2}{\ignorespaces Performance Comparison of Surrogate Models for Gear Deviation Prediction}}{16}{}\protected@file@percent }
\newlabel{tab:ml_comparison}{{2}{16}}
\@writefile{toc}{\contentsline {subsubsection}{\numberline {4.1.1}Analysis of Model Performance}{16}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {4.2}Ablation Study and Performance Analysis}{17}{}\protected@file@percent }
\@writefile{lof}{\contentsline {figure}{\numberline {5}{\ignorespaces Performance trend of different architectures with cumulative enhancements. While DQN (blue) and Double DQN (red) exhibit significant volatility, the proposed Residual Dueling DQN (green) demonstrates robust and monotonic improvement, ultimately achieving the highest success rate of 92\%.}}{18}{}\protected@file@percent }
\newlabel{fig:ablation_trend}{{5}{18}}
\@writefile{lot}{\contentsline {table}{\numberline {3}{\ignorespaces Comparative Performance Analysis of DQN, Double DQN, and Residual Dueling DQN. The table details the success counts, step distributions, average steps, and final positioning error (RMSE) over 25 test episodes.}}{18}{}\protected@file@percent }
\newlabel{tab:model_comparison}{{3}{18}}
\@writefile{toc}{\contentsline {subsubsection}{\numberline {4.2.1}Analysis of Ablation Study}{19}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsubsection}{\numberline {4.2.2}Comparative Analysis of Final Models}{19}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {4.3}Visualization of Experimental Results}{20}{}\protected@file@percent }
\@writefile{lof}{\contentsline {figure}{\numberline {6}{\ignorespaces Visualization of sequential optimization trajectories across representative test scenarios. The purple star denotes the global minimum, and the green line traces the agent's decision path. The horizontal axis represents the physical one-dimensional shim displacement, while the vertical axis plots a composite geometric deviation profile obtained by taking the maximum envelope across the eight transmission-error metrics. Although the RL agent is trained using the complete 13-dimensional state, this figure visualizes the projected one-dimensional optimization landscape derived from the eight gear-deviation descriptors. Subfigure (a) depicts successful convergence to the target zone, while (b) illustrates a localized precision miss under non-convex topology challenges.}}{20}{}\protected@file@percent }
\newlabel{fig:search_trajectories}{{6}{20}}
\@writefile{toc}{\contentsline {subsection}{\numberline {4.4}Overall Efficiency and Industrial Decision Value}{20}{}\protected@file@percent }
\citation{Feng2020,Srivastava2019}
\citation{wu2022novel}
\@writefile{toc}{\contentsline {section}{\numberline {5}Discussion}{21}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {5.1}Physical Interpretability of Random Forest as a Virtual Environment}{21}{}\protected@file@percent }
\citation{Wang2016,Liu2025}
\@writefile{toc}{\contentsline {subsection}{\numberline {5.2}Analysis of the RL Convergence Mechanism}{22}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {5.3}Methodological Generalizability and Framework Scalability}{22}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {5.4}Limitations}{22}{}\protected@file@percent }
\@writefile{toc}{\contentsline {section}{\numberline {6}Conclusions and Future Work}{23}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {6.1}Summary of Findings}{23}{}\protected@file@percent }
\citation{kuo2025deep}
\citation{cao2024digital}
\citation{demvcak2024digital}
\@writefile{toc}{\contentsline {subsection}{\numberline {6.2}Research Significance and Industrial Value}{24}{}\protected@file@percent }
\@writefile{toc}{\contentsline {subsection}{\numberline {6.3}Future Work}{24}{}\protected@file@percent }
\bibstyle{unsrt}
\bibdata{References}
\bibcite{Luo2017}{{1}{}{{}}{{}}}
\bibcite{Ding2025}{{2}{}{{}}{{}}}
\bibcite{Deng2022}{{3}{}{{}}{{}}}
\bibcite{Paranjape2023}{{4}{}{{}}{{}}}
\bibcite{Zhu2023}{{5}{}{{}}{{}}}
\bibcite{Ming2023}{{6}{}{{}}{{}}}
\bibcite{Li2024}{{7}{}{{}}{{}}}
\bibcite{Sadeghian2025}{{8}{}{{}}{{}}}
\bibcite{Jiang2024}{{9}{}{{}}{{}}}
\bibcite{Jin2011}{{10}{}{{}}{{}}}
\bibcite{AGMA2015}{{11}{}{{}}{{}}}
\bibcite{Wang2016}{{12}{}{{}}{{}}}
\bibcite{He2016}{{13}{}{{}}{{}}}
\bibcite{Hester2018}{{14}{}{{}}{{}}}
\bibcite{ElYousfi2021}{{15}{}{{}}{{}}}
\bibcite{Rao2020}{{16}{}{{}}{{}}}
\bibcite{Wu2023}{{17}{}{{}}{{}}}
\bibcite{Xu2018}{{18}{}{{}}{{}}}
\bibcite{Chhor2021}{{19}{}{{}}{{}}}
\bibcite{Feng2020}{{20}{}{{}}{{}}}
\bibcite{Tsang2025}{{21}{}{{}}{{}}}
\bibcite{wrede2024}{{22}{}{{}}{{}}}
\bibcite{espeholt2018impala}{{23}{}{{}}{{}}}
\bibcite{machado2018revisiting}{{24}{}{{}}{{}}}
\bibcite{Schaul2015}{{25}{}{{}}{{}}}
\bibcite{Srivastava2019}{{26}{}{{}}{{}}}
\bibcite{wu2022novel}{{27}{}{{}}{{}}}
\bibcite{Liu2025}{{28}{}{{}}{{}}}
\bibcite{kuo2025deep}{{29}{}{{}}{{}}}
\bibcite{cao2024digital}{{30}{}{{}}{{}}}
\bibcite{demvcak2024digital}{{31}{}{{}}{{}}}
\providecommand\NAT@force@numbers{}\NAT@force@numbers
\gdef \@abspage@last{28}
