@misc{indiciae9a623f4f88cb, title = {Activating Visual Context and Commonsense Reasoning through Masked Prediction in VLMs}, author = {Jiaao Yu and Shenwei Li and Mingjie Han and Yifei Yin and Wenzheng Song and Chenghao Jia and Man Lan}, year = {2025}, url = {https://arxiv.org/abs/2510.21807}, note = {Source identifier: 2510.21807} }