@misc{indiciae858c000abdde, title = {Teach a Reward Model to Correct Itself: Reward Guided Adversarial Failure Discovery for Robust Reward Modeling}, author = {Pankayaraj Pathmanathan and Furong Huang}, year = {2026}, url = {https://arxiv.org/abs/2507.06419}, note = {Source identifier: 2507.06419} }