@misc{indiciae82f7d06eb349, title = {TernaryCLIP: Efficiently Compressing Vision-Language Models with Ternary Weights and Distilled Knowledge}, author = {Shu-Hao Zhang and Wei-Cheng Tang and Chen Wu and Peng Hu and Nan Li and Liang-Jie Zhang and Qi Zhang and Shao-Qun Zhang}, year = {2025}, url = {https://arxiv.org/abs/2510.21879}, note = {Source identifier: 2510.21879} }