@misc{indiciae75161700a0f3, title = {Gradient Regularization Mitigates Reward Hacking in Reinforcement Learning from Human Feedback and Verifiable Rewards}, author = {Johannes Ackermann and Michael Noukhovitch and Takashi Ishida and Masashi Sugiyama}, year = {2026}, url = {https://arxiv.org/abs/2602.18037}, note = {Source identifier: 2602.18037} }