<?xml version="1.0" encoding="UTF-8"?><ns2:project xmlns:ns1="http://gtr.rcuk.ac.uk/gtr/api" xmlns:ns2="http://gtr.rcuk.ac.uk/gtr/api/project" xmlns:ns3="http://gtr.rcuk.ac.uk/gtr/api/fund" xmlns:ns4="http://gtr.rcuk.ac.uk/gtr/api/person" xmlns:ns5="http://gtr.rcuk.ac.uk/gtr/api/project/outcome" xmlns:ns6="http://gtr.rcuk.ac.uk/gtr/api/organisation" ns1:created="2026-07-08T08:44:08Z" ns1:href="http://gtr.ukri.org/gtr/api/projects/208B8687-DC7C-4E15-8DD6-E3AB6834E1EF" ns1:id="208B8687-DC7C-4E15-8DD6-E3AB6834E1EF"><ns1:links><ns1:link ns1:href="http://gtr.ukri.org/gtr/api/persons/7646C46D-24A6-4227-B47A-47692E7BA500" ns1:rel="PM_PER"/><ns1:link ns1:href="http://gtr.ukri.org/gtr/api/organisations/4A0F7975-E20E-40BF-BB53-5F035E24415F" ns1:rel="LEAD_ORG"/><ns1:link ns1:href="http://gtr.ukri.org/gtr/api/organisations/4A0F7975-E20E-40BF-BB53-5F035E24415F" ns1:rel="PARTICIPANT_ORG"/><ns1:link ns1:href="http://gtr.ukri.org/gtr/api/organisations/6D906590-BA40-4EFA-A713-A8638A0BA948" ns1:rel="PARTICIPANT_ORG"/><ns1:link ns1:end="2012-01-31T00:00:00Z" ns1:href="http://gtr.ukri.org/gtr/api/funds/9030A0B9-2586-4F2B-B1A5-16283C05FE0B" ns1:rel="FUND" ns1:start="2011-11-01T00:00:00Z"/></ns1:links><ns2:identifiers><ns2:identifier ns2:type="RCUK">130739</ns2:identifier></ns2:identifiers><ns2:title>Building on subtitles as a source of live metadata for UK broadcast TV</ns2:title><ns2:status>Closed</ns2:status><ns2:grantCategory>Fast Track</ns2:grantCategory><ns2:leadFunder>Innovate UK</ns2:leadFunder><ns2:abstractText>TV Subtitles are already a rich source of pre-existing metadata. This has been exploited by some to provide advanced search techniques for video archives such ashttp://www.hulu.com/labs/captions-search.
Modern Natural Language Processing (NLP) techniques make it possible to extract more detailed and structured information from unstructured scripts. This includes indentifying people, places, political concepts etc.
We propose to create a system that will extract subtitles from the broadcast stream and feed them through a concept extraction engine in real time to create a new layer of semantic metadata</ns2:abstractText></ns2:project>