@misc{indiciae76bc9f4f9400, title = {Text encoders bottleneck compositionality in contrastive vision-language models}, author = {Amita Kamath and Jack Hessel and Kai-Wei Chang}, year = {2023}, url = {https://arxiv.org/abs/2305.14897}, note = {Source identifier: 2305.14897} }