@misc{indiciae55784143dab8, title = {Causally Robust Reward Learning from Reason-Augmented Preference Feedback}, author = {Minjune Hwang and Yigit Korkmaz and Daniel Seita and Erdem Bıyık}, year = {2026}, url = {https://arxiv.org/abs/2603.04861}, note = {Source identifier: 2603.04861} }