@misc{indiciae8009398b1b38, title = {Aligning What Vision-Language Models See and Perceive with Adaptive Information Flow}, author = {Chengxin Liu and Wonseok Choi and Chenshuang Zhang and Tae-Hyun Oh}, year = {2026}, url = {https://arxiv.org/abs/2604.15809}, note = {Source identifier: 2604.15809} }