@misc{indiciae13e68a97c8e9, title = {Reason, Reward, Refine: Step-Level Errors Corrections with Structured Feedback for Physics Reasoning in Small Language Models}, author = {Raj Jaiswal and Dhruv Jain and Rishabh Dhawan and Sree Krishna Uppalapati and Shin'ichi Satoh and Tanuja Ganu and Rajiv Ratn Shah}, year = {2026}, url = {https://arxiv.org/abs/2607.05199}, note = {Source identifier: 2607.05199} }