@misc{indiciaea554522706ff, title = {Generalization Limits of Reinforcement Learning Alignment}, author = {Haruhi Shida and Koo Imai and Keigo Kansa}, year = {2026}, url = {https://arxiv.org/abs/2604.02652}, note = {Source identifier: 2604.02652} }