@misc{indiciae5ec6db2e814b, title = {OmniVLM: A Token-Compressed, Sub-Billion-Parameter Vision-Language Model for Efficient On-Device Inference}, author = {Wei Chen and Zhiyuan Li and Shuo Xin}, year = {2026}, url = {https://arxiv.org/abs/2412.11475}, note = {Source identifier: 2412.11475} }