@misc{indiciae2513ac5a624b, title = {SmartCLIP: Modular Vision-language Alignment with Identification Guarantees}, author = {Shaoan Xie and Lingjing Kong and Yujia Zheng and Yu Yao and Zeyu Tang and Eric P. Xing and Guangyi Chen and Kun Zhang}, year = {2026}, url = {https://arxiv.org/abs/2507.22264}, note = {Source identifier: 2507.22264} }