@misc{indiciae5afc08c6160d, title = {Focus-then-Context: Subject-Centric Progressive Visual Token Reduction for Vision-Language Models}, author = {Yulin Zhao and Zheng Zhang}, year = {2026}, url = {https://arxiv.org/abs/2605.20950}, note = {Source identifier: 2605.20950} }