@misc{indiciae08c97b038fa4, title = {Can Large Language Models Grasp Concepts in Visual Content? A Case Study on YouTube Shorts about Depression}, author = {Jiaying "Lizzy" Liu and Yiheng Su and Praneel Seth}, year = {2025}, doi = {10.1145/3706599.3719821}, url = {https://arxiv.org/abs/2503.05109}, note = {Source identifier: 2503.05109} }