@misc{indiciae99504d9ed252, title = {Learning Domain Knowledge in Multimodal Large Language Models through Reinforcement Fine-Tuning}, author = {Qinglong Cao and Yuntian Chen and Chao Ma and Xiaokang Yang}, year = {2026}, url = {https://arxiv.org/abs/2601.16419}, note = {Source identifier: 2601.16419} }