youtube-dl/youtube_dl/extractor/movieclips.py

# coding: utf-8
from __future__ import unicode_literals

from .common import InfoExtractor
from ..utils import (
    smuggle_url,
    float_or_none,
    parse_iso8601,
    update_url_query,
)


class MovieClipsIE(InfoExtractor):
    _VALID_URL = r'https?://(?:www\.)?movieclips\.com/videos/.+-(?P<id>\d+)(?:\?|$)'
    _TEST = {
        'url': 'http://www.movieclips.com/videos/warcraft-trailer-1-561180739597',
        'md5': '42b5a0352d4933a7bd54f2104f481244',
        'info_dict': {
            'id': 'pKIGmG83AqD9',
            'ext': 'mp4',
            'title': 'Warcraft Trailer 1',
            'description': 'Watch Trailer 1 from Warcraft (2016). Legendary’s WARCRAFT is a 3D epic adventure of world-colliding conflict based.',
            'thumbnail': r're:^https?://.*\.jpg$',
            'timestamp': 1446843055,
            'upload_date': '20151106',
            'uploader': 'Movieclips',
        },
        'add_ie': ['ThePlatform'],
    }

    def _real_extract(self, url):
        video_id = self._match_id(url)
        webpage = self._download_webpage(url, video_id)
        video = next(v for v in self._parse_json(self._search_regex(
            r'var\s+__REACT_ENGINE__\s*=\s*({.+});',
            webpage, 'react engine'), video_id)['playlist']['videos'] if v['id'] == video_id)

        return {
            '_type': 'url_transparent',
            'ie_key': 'ThePlatform',
            'url': smuggle_url(update_url_query(
                video['contentUrl'], {'mbr': 'true'}), {'force_smil_url': True}),
            'title': self._og_search_title(webpage),
            'description': self._html_search_meta('description', webpage),
            'duration': float_or_none(video.get('duration')),
            'timestamp': parse_iso8601(video.get('dateCreated')),
            'thumbnail': video.get('defaultImage'),
            'uploader': video.get('provider'),
        }
-												[movieclips] Add coding cookie

											
										
										
											2015-11-08 11:56:20 +01:00
+								# coding: utf-8
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
+								from __future__ import unicode_literals
 								from .common import InfoExtractor
-												[movieclips] fix extraction

											
										
										
											2016-04-01 17:22:06 +02:00
+								from ..utils import (
 								    smuggle_url,
 								    float_or_none,
 								    parse_iso8601,
 								    update_url_query,
 								)
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
 								class MovieClipsIE(InfoExtractor):
-												[movieclips] Fix _VALID_URL

											
										
										
											2016-10-24 18:54:03 +02:00
+								    _VALID_URL = r'https?://(?:www\.)?movieclips\.com/videos/.+-(?P<id>\d+)(?:\?|$)'
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
+								    _TEST = {
-												[movieclips] fix extraction

											
										
										
											2016-04-01 17:22:06 +02:00
+								        'url': 'http://www.movieclips.com/videos/warcraft-trailer-1-561180739597',
 								        'md5': '42b5a0352d4933a7bd54f2104f481244',
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
+								        'info_dict': {
-												[movieclips] Fix extraction (fixes #7404)

They use theplatform now.
Changed the test, because the old one seems to be georestricted.

											
										
										
											2015-11-08 11:49:51 +01:00
+								            'id': 'pKIGmG83AqD9',
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
+								            'ext': 'mp4',
-												[movieclips] Fix extraction (fixes #7404)

They use theplatform now.
Changed the test, because the old one seems to be georestricted.

											
										
										
											2015-11-08 11:49:51 +01:00
+								            'title': 'Warcraft Trailer 1',
 								            'description': 'Watch Trailer 1 from Warcraft (2016). Legendary’s WARCRAFT is a 3D epic adventure of world-colliding conflict based.',
-												Fix "invalid escape sequences" error on Python 3.6

											
										
										
											2017-01-02 13:08:07 +01:00
+								            'thumbnail': r're:^https?://.*\.jpg$',
-												[movieclips] fix extraction

											
										
										
											2016-04-01 17:22:06 +02:00
+								            'timestamp': 1446843055,
 								            'upload_date': '20151106',
 								            'uploader': 'Movieclips',
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
+								        },
-												[movieclips] Fix extraction (fixes #7404)

They use theplatform now.
Changed the test, because the old one seems to be georestricted.

											
										
										
											2015-11-08 11:49:51 +01:00
+								        'add_ie': ['ThePlatform'],
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
+								    }
 								    def _real_extract(self, url):
-												[movieclips] fix extraction

											
										
										
											2016-04-01 17:22:06 +02:00
+								        video_id = self._match_id(url)
 								        webpage = self._download_webpage(url, video_id)
 								        video = next(v for v in self._parse_json(self._search_regex(
 								            r'var\s+__REACT_ENGINE__\s*=\s*({.+});',
 								            webpage, 'react engine'), video_id)['playlist']['videos'] if v['id'] == video_id)
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
 								        return {
-												[movieclips] Fix extraction (fixes #7404)

They use theplatform now.
Changed the test, because the old one seems to be georestricted.

											
										
										
											2015-11-08 11:49:51 +01:00
+								            '_type': 'url_transparent',
-												[movieclips] fix extraction

											
										
										
											2016-04-01 17:22:06 +02:00
+								            'ie_key': 'ThePlatform',
 								            'url': smuggle_url(update_url_query(
 								                video['contentUrl'], {'mbr': 'true'}), {'force_smil_url': True}),
 								            'title': self._og_search_title(webpage),
 								            'description': self._html_search_meta('description', webpage),
 								            'duration': float_or_none(video.get('duration')),
 								            'timestamp': parse_iso8601(video.get('dateCreated')),
 								            'thumbnail': video.get('defaultImage'),
 								            'uploader': video.get('provider'),
-												[movieclips] Add extractor (Closes #3554)

											
										
										
											2014-08-23 12:44:56 +02:00
+								        }