@misc{indiciaeeadd0184af02, title = {RGBX-R1: Visual Modality Chain-of-Thought Guided Reinforcement Learning for Multimodal Grounding}, author = {Jiahe Wu and Bing Cao and Qilong Wang and Qinghua Hu and Dongdong Li and Pengfei Zhu}, year = {2026}, url = {https://arxiv.org/abs/2602.00504}, note = {Source identifier: 2602.00504} }