@misc{indiciae3e7f379ce4a4, title = {Attend, Transform, or Silence: Operator-Level Visual Skipping for Efficient Multimodal LLM Inference}, author = {Zhaoyang Luo and Runmin Dong and Miao Yang and Fan Wei and Yushan Lai and Bin Luo and Haohuan Fu}, year = {2026}, url = {https://arxiv.org/abs/2606.31903}, note = {Source identifier: 2606.31903} }