@misc{indiciaef812555208bb, title = {Understanding Contexts Inside Robot and Human Manipulation Tasks through a Vision-Language Model and Ontology System in a Video Stream}, author = {Chen Jiang and Masood Dehghan and Martin Jagersand}, year = {2020}, url = {https://arxiv.org/abs/2003.01163}, note = {Source identifier: 2003.01163} }