@misc{indiciaea49a41ad319f, title = {AdaptVision: Dynamic Input Scaling in MLLMs for Versatile Scene Understanding}, author = {Yonghui Wang and Wengang Zhou and Hao Feng and Houqiang Li}, year = {2024}, url = {https://arxiv.org/abs/2408.16986}, note = {Source identifier: 2408.16986} }