@misc{indiciaee79a640b5b6c, title = {LiteVLM: A Low-Latency Vision-Language Model Inference Pipeline for Resource-Constrained Environments}, author = {Jin Huang and Yuchao Jin and Le An and Josh Park}, year = {2025}, url = {https://arxiv.org/abs/2506.07416}, note = {Source identifier: 2506.07416} }