@misc{indiciae4f4872d05ed2, title = {Detecting Data Contamination from Reinforcement Learning Post-training for Large Language Models}, author = {Yongding Tao and Tian Wang and Yihong Dong and Huanyu Liu and Kechi Zhang and Xiaolong Hu and Ge Li}, year = {2026}, url = {https://arxiv.org/abs/2510.09259}, note = {Source identifier: 2510.09259} }