@inproceedings{azad2026streamready,title={StreamReady: Learning What to Answer and When in Long Streaming Videos},author={Azad, Shehreen and Vineet, Vibhav and Rawat, Yogesh},booktitle={IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)},pages={40494--40504},year={2026},}
Forget, Anticipate and Adapt: Test Time Training for Long Videos
Rajat Modi, S. Noel, Xin Liang, and Yogesh S. Rawat
In European Conference on Computer Vision (ECCV), 2026
@inproceedings{modi2026forget,title={Forget, Anticipate and Adapt: Test Time Training for Long Videos},author={Modi, Rajat and Noel, S. and Liang, Xin and Rawat, Yogesh S.},booktitle={European Conference on Computer Vision (ECCV)},year={2026},}
2025
HierarQ: Task-Aware Hierarchical Q-Former for Enhanced Video Understanding
Shehreen Azad, Vibhav Vineet, and Yogesh Singh Rawat
In IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR), 2025
@inproceedings{azad2025hierarq,title={HierarQ: Task-Aware Hierarchical Q-Former for Enhanced Video Understanding},author={Azad, Shehreen and Vineet, Vibhav and Rawat, Yogesh Singh},booktitle={IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)},pages={8545--8556},year={2025},}
A Large-Scale Analysis on Contextual Self-Supervised Video Representation Learning
Akash Kumar, Ashlesha Kumar, Vibhav Vineet, and Yogesh S Rawat
In Proceedings of the Computer Vision and Pattern Recognition Conference, 2025
@inproceedings{kumar2025large,title={A Large-Scale Analysis on Contextual Self-Supervised Video Representation Learning},author={Kumar, Akash and Kumar, Ashlesha and Vineet, Vibhav and Rawat, Yogesh S},booktitle={Proceedings of the Computer Vision and Pattern Recognition Conference},pages={670--681},year={2025},}
2024
Foundation models for video understanding: A survey
Neelu Madan, Andreas Møgelmose, Rajat Modi, Yogesh S Rawat, and Thomas B Moeslund
@article{madan2024foundation,title={Foundation models for video understanding: A survey},author={Madan, Neelu and M{\o}gelmose, Andreas and Modi, Rajat and Rawat, Yogesh S and Moeslund, Thomas B},journal={arXiv preprint arXiv:2405.03770},year={2024},}
2023
Self-supervised learning for videos: A survey
Madeline C Schiappa, Yogesh S Rawat, and Mubarak Shah
@article{schiappa2023self,title={Self-supervised learning for videos: A survey},author={Schiappa, Madeline C and Rawat, Yogesh S and Shah, Mubarak},journal={ACM Computing Surveys},volume={55},number={13s},pages={1--37},year={2023},publisher={ACM New York, NY},}
2022
Svgraph: Learning semantic graphs from instructional videos
Madeline C Schiappa and Yogesh S Rawat
2022 IEEE Eighth International Conference on Multimedia Big Data (BigMM), 2022
@article{schiappa2022svgraph,title={Svgraph: Learning semantic graphs from instructional videos},author={Schiappa, Madeline C and Rawat, Yogesh S},journal={2022 IEEE Eighth International Conference on Multimedia Big Data (BigMM)},pages={45--52},year={2022},}