@Article{JordiGonzalez2009, author="Jordi Gonzalez and Dani Rowe and J. Varona and Xavier Roca", title="Understanding Dynamic Scenes based on Human Sequence Evaluation", journal="Image and Vision Computing", year="2009", volume="27", number="10", pages="1433--1444", optkeywords="Image Sequence Evaluation", optkeywords="High-level processing of monitored scenes", optkeywords="Segmentation and tracking in complex scenes", optkeywords="Event recognition in dynamic scenes", optkeywords="Human motion understanding", optkeywords="Human behaviour interpretation", optkeywords="Natural-language text generation", optkeywords="Realistic demonstrators", abstract="In this paper, a Cognitive Vision System (CVS) is presented, which explains the human behaviour of monitored scenes using natural-language texts. This cognitive analysis of human movements recorded in image sequences is here referred to as Human Sequence Evaluation (HSE) which defines a set of transformation modules involved in the automatic generation of semantic descriptions from pixel values. In essence, the trajectories of human agents are obtained to generate textual interpretations of their motion, and also to infer the conceptual relationships of each agent w.r.t. its environment. For this purpose, a human behaviour model based on Situation Graph Trees (SGTs) is considered, which permits both bottom-up (hypothesis generation) and top-down (hypothesis refinement) analysis of dynamic scenes. The resulting system prototype interprets different kinds of behaviour and reports textual descriptions in multiple languages.", optnote="ISE", optnote="exported from refbase (http://refbase.cvc.uab.es/show.php?record=1211), last updated on Thu, 08 Sep 2016 18:39:54 +0200", doi="10.1016/j.imavis.2008.02.004" }