@misc{indiciaea537b3b97205, title = {DeCap: Decoding CLIP Latents for Zero-Shot Captioning via Text-Only Training}, author = {Wei Li and Linchao Zhu and Longyin Wen and Yi Yang}, year = {2023}, url = {https://arxiv.org/abs/2303.03032}, note = {Source identifier: 2303.03032} }