@inproceedings{tvqaspatiotemporalgroundingforvideo, title = {TVQA+: Spatio-Temporal Grounding for Video Question Answering}, author = {Jie Lei and Licheng Yu and Tamara L. Berg and Mohit Bansal}, year = {2019}, booktitle = {ACL 2020 6}, eprint = {1904.11574}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/1904.11574v2}, }