@misc{indiciaee57d9230d435, title = {J1: Incentivizing Thinking in LLM-as-a-Judge via Reinforcement Learning}, author = {Chenxi Whitehouse and Tianlu Wang and Ping Yu and Xian Li and Jason Weston and Ilia Kulikov and Swarnadeep Saha}, year = {2025}, url = {https://arxiv.org/abs/2505.10320}, note = {Source identifier: 2505.10320} }