youtube-dl/youtube_dl/extractor/spiegeltv.py

# coding: utf-8
from __future__ import unicode_literals

import re
from .common import InfoExtractor


class SpiegeltvIE(InfoExtractor):
    _VALID_URL = r'https?://(?:www\.)?spiegel\.tv/filme/(?P<id>[\-a-z0-9]+)'
    _TEST = {
        'url': 'http://www.spiegel.tv/filme/flug-mh370/',
        'info_dict': {
            'id': 'flug-mh370',
            'ext': 'm4v',
            'title': 'Flug MH370',
            'description': 'Das Rätsel um die Boeing 777 der Malaysia-Airlines',
            'thumbnail': 're:http://.*\.jpg$',
        },
        'params': {
            # rtmp download
            'skip_download': True,
        }
    }

    def _real_extract(self, url):
        mobj = re.match(self._VALID_URL, url)
        video_id = mobj.group('id')

        webpage = self._download_webpage(url, video_id)
        title = self._html_search_regex(r'<h1.*?>(.*?)</h1>', webpage, 'title')

        apihost = 'http://spiegeltv-ivms2-restapi.s3.amazonaws.com'
        version_json = self._download_json(
            '%s/version.json' % apihost, video_id,
            note='Downloading version information')
        version_name = version_json['version_name']

        slug_json = self._download_json(
            '%s/%s/restapi/slugs/%s.json' % (apihost, version_name, video_id),
            video_id,
            note='Downloading object information')
        oid = slug_json['object_id']

        media_json = self._download_json(
            '%s/%s/restapi/media/%s.json' % (apihost, version_name, oid),
            video_id, note='Downloading media information')
        uuid = media_json['uuid']
        is_wide = media_json['is_wide']

        server_json = self._download_json(
            'http://www.spiegel.tv/streaming_servers/', video_id,
            note='Downloading server information')
        server = server_json[0]['endpoint']

        thumbnails = []
        for image in media_json['images']:
            thumbnails.append({
                'url': image['url'],
                'width': image['width'],
                'height': image['height'],
            })

        description = media_json['subtitle']
        duration = media_json['duration_in_ms'] / 1000.

        if is_wide:
            format = '16x9'
        else:
            format = '4x3'

        url = server + 'mp4:' + uuid + '_spiegeltv_0500_' + format + '.m4v'

        return {
            'id': video_id,
            'title': title,
            'url': url,
            'ext': 'm4v',
            'description': description,
            'duration': duration,
            'thumbnails': thumbnails
        }
added spiegel.tv 2014-05-30 22:35:17 +08:00			`# coding: utf-8`
			`from __future__ import unicode_literals`

			`import re`
			`from .common import InfoExtractor`

[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00
added spiegel.tv 2014-05-30 22:35:17 +08:00			`class SpiegeltvIE(InfoExtractor):`
			`_VALID_URL = r'https?://(?:www\.)?spiegel\.tv/filme/(?P<id>[\-a-z0-9]+)'`
			`_TEST = {`
			`'url': 'http://www.spiegel.tv/filme/flug-mh370/',`
			`'info_dict': {`
			`'id': 'flug-mh370',`
			`'ext': 'm4v',`
			`'title': 'Flug MH370',`
			`'description': 'Das Rätsel um die Boeing 777 der Malaysia-Airlines',`
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`'thumbnail': 're:http://.*\.jpg$',`
[Spiegeltv] skip rtmp download to pass Travis test build 2014-06-03 22:50:54 +08:00			`},`
			`'params': {`
			`# rtmp download`
			`'skip_download': True,`
added spiegel.tv 2014-05-30 22:35:17 +08:00			`}`
			`}`

			`def _real_extract(self, url):`
			`mobj = re.match(self._VALID_URL, url)`
			`video_id = mobj.group('id')`

			`webpage = self._download_webpage(url, video_id)`
			`title = self._html_search_regex(r'<h1.?>(.?)</h1>', webpage, 'title')`

[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`apihost = 'http://spiegeltv-ivms2-restapi.s3.amazonaws.com'`
			`version_json = self._download_json(`
			`'%s/version.json' % apihost, video_id,`
			`note='Downloading version information')`
			`version_name = version_json['version_name']`
added spiegel.tv 2014-05-30 22:35:17 +08:00
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`slug_json = self._download_json(`
			`'%s/%s/restapi/slugs/%s.json' % (apihost, version_name, video_id),`
			`video_id,`
			`note='Downloading object information')`
			`oid = slug_json['object_id']`
added spiegel.tv 2014-05-30 22:35:17 +08:00
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`media_json = self._download_json(`
			`'%s/%s/restapi/media/%s.json' % (apihost, version_name, oid),`
			`video_id, note='Downloading media information')`
			`uuid = media_json['uuid']`
			`is_wide = media_json['is_wide']`
added spiegel.tv 2014-05-30 22:35:17 +08:00
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`server_json = self._download_json(`
			`'http://www.spiegel.tv/streaming_servers/', video_id,`
			`note='Downloading server information')`
			`server = server_json[0]['endpoint']`
added spiegel.tv 2014-05-30 22:35:17 +08:00
			`thumbnails = []`
			`for image in media_json['images']:`
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`thumbnails.append({`
			`'url': image['url'],`
			`'width': image['width'],`
			`'height': image['height'],`
			`})`
added spiegel.tv 2014-05-30 22:35:17 +08:00
			`description = media_json['subtitle']`
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`duration = media_json['duration_in_ms'] / 1000.`
added spiegel.tv 2014-05-30 22:35:17 +08:00
			`if is_wide:`
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`format = '16x9'`
added spiegel.tv 2014-05-30 22:35:17 +08:00			`else:`
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`format = '4x3'`
added spiegel.tv 2014-05-30 22:35:17 +08:00
			`url = server + 'mp4:' + uuid + '_spiegeltv_0500_' + format + '.m4v'`

[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`return {`
added spiegel.tv 2014-05-30 22:35:17 +08:00			`'id': video_id,`
			`'title': title,`
			`'url': url,`
			`'ext': 'm4v',`
			`'description': description,`
			`'duration': duration,`
			`'thumbnails': thumbnails`
[spiegeltv] Simplify and PEP8 2014-06-07 21:33:45 +08:00			`}`