@misc{indiciaee1caf0c667b8, title = {IR\$\textasciicircum{}3\$: Contrastive Inverse Reinforcement Learning for Interpretable Detection and Mitigation of Reward Hacking}, author = {Mohammad Beigi and Ming Jin and Junshan Zhang and Jiaxin Zhang and Qifan Wang and Lifu Huang}, year = {2026}, url = {https://arxiv.org/abs/2602.19416}, note = {Source identifier: 2602.19416} }