@misc{indiciae2a49aa2ab0fb, title = {Docopilot: Improving Multimodal Models for Document-Level Understanding}, author = {Yuchen Duan and Zhe Chen and Yusong Hu and Weiyun Wang and Shenglong Ye and Botian Shi and Lewei Lu and Qibin Hou and Tong Lu and Hongsheng Li and Jifeng Dai and Wenhai Wang}, year = {2025}, url = {https://arxiv.org/abs/2507.14675}, note = {Source identifier: 2507.14675} }