@misc{indiciaeeb4c61a6a4c4, title = {AddressCLIP: Empowering Vision-Language Models for City-wide Image Address Localization}, author = {Shixiong Xu and Chenghao Zhang and Lubin Fan and Gaofeng Meng and Shiming Xiang and Jieping Ye}, year = {2024}, url = {https://arxiv.org/abs/2407.08156}, note = {Source identifier: 2407.08156} }