@misc{indiciae7caccff093d7, title = {ViDscribe: Multimodal AI for Customizing Audio Description and Question Answering in Online Videos}, author = {Maryam Cheema and Sina Elahimanesh and Pooyan Fazli and Hasti Seifi}, year = {2026}, doi = {10.1145/3772363.3798744}, url = {https://arxiv.org/abs/2603.14662}, note = {Source identifier: 2603.14662} }