@misc{indiciae8fd48846da35, title = {Directional Alignment Mitigates Reward Hacking in Reinforcement Learning for Language Models}, author = {Wenlong Deng and Jiaji Huang and Kaan Ozkara and Yushu Li and Christos Thrampoulidis and Xiaoxiao Li and Youngsuk Park}, year = {2026}, url = {https://arxiv.org/abs/2605.25189}, note = {Source identifier: 2605.25189} }