@misc{indiciaefd89320198ec, title = {Attention-Guided Reward for Reinforcement Learning-based Jailbreak against Large Reasoning Models}, author = {Zheng Lin and Zhenxing Niu and Haoxuan Ji and Yuzhe Huang and Haichang Gao}, year = {2026}, url = {https://arxiv.org/abs/2605.19485}, note = {Source identifier: 2605.19485} }