@misc{indiciae5c3db0218fe6, title = {Reinforcement Learning Tuning for VideoLLMs: Reward Design and Data Efficiency}, author = {Hongyu Li and Songhao Han and Yue Liao and Junfeng Luo and Jialin Gao and Shuicheng Yan and Si Liu}, year = {2025}, url = {https://arxiv.org/abs/2506.01908}, note = {Source identifier: 2506.01908} }