@misc{indiciaecf4d95045686, title = {Enhancing Audio-Language Models through Self-Supervised Post-Training with Text-Audio Pairs}, author = {Anshuman Sinha and Camille Migozzi and Aubin Rey and Chao Zhang}, year = {2025}, url = {https://arxiv.org/abs/2408.09269}, note = {Source identifier: 2408.09269} }