@misc{indiciae872b5ce66842, title = {Leveraging Vision-Language Pre-training for Human Activity Recognition in Still Images}, author = {Cristina Mahanta and Gagan Bhatia}, year = {2025}, url = {https://arxiv.org/abs/2506.13458}, note = {Source identifier: 2506.13458} }