@misc{indiciaef6dca9a98260, title = {VL-Rethinker: Incentivizing Self-Reflection of Vision-Language Models with Reinforcement Learning}, author = {Haozhe Wang and Chao Qu and Zuming Huang and Wei Chu and Fangzhen Lin and Wenhu Chen}, year = {2025}, url = {https://arxiv.org/abs/2504.08837}, note = {Source identifier: 2504.08837} }