@misc{indiciae3d1f9bf15bde, title = {Cross-Modal Concept Learning and Inference for Vision-Language Models}, author = {Yi Zhang and Ce Zhang and Yushun Tang and Zhihai He}, year = {2023}, url = {https://arxiv.org/abs/2307.15460}, note = {Source identifier: 2307.15460} }