@misc{indiciae9820e90bc0c1, title = {V2P-Bench: Evaluating Video-Language Understanding with Visual Prompts for Better Human-Model Interaction}, author = {Yiming Zhao and Yu Zeng and Yukun Qi and YaoYang Liu and Xikun Bao and Lin Chen and Zehui Chen and Qing Miao and Chenxi Liu and Jie Zhao and Feng Zhao}, year = {2026}, url = {https://arxiv.org/abs/2503.17736}, note = {Source identifier: 2503.17736} }