@misc{indiciae9035e0d7aa84, title = {Reading Images Like Texts: Sequential Image Understanding in Vision-Language Models}, author = {Yueyan Li and Chenggong Zhao and Zeyuan Zang and Caixia Yuan and Xiaojie Wang}, year = {2026}, url = {https://arxiv.org/abs/2509.19191}, note = {Source identifier: 2509.19191} }