diff options
Diffstat (limited to 'youtube_dl/extractor/arte.py')
| -rw-r--r-- | youtube_dl/extractor/arte.py | 39 | 
1 files changed, 37 insertions, 2 deletions
diff --git a/youtube_dl/extractor/arte.py b/youtube_dl/extractor/arte.py index 8b62ee774..4b7bef775 100644 --- a/youtube_dl/extractor/arte.py +++ b/youtube_dl/extractor/arte.py @@ -10,6 +10,7 @@ from ..utils import (      determine_ext,      get_element_by_id,      compat_str, +    get_element_by_attribute,  )  # There are different sources of video in arte.tv, the extraction process  @@ -17,8 +18,8 @@ from ..utils import (  # add tests.  class ArteTvIE(InfoExtractor): -    _VIDEOS_URL = r'(?:http://)?videos.arte.tv/(?P<lang>fr|de)/.*-(?P<id>.*?).html' -    _LIVEWEB_URL = r'(?:http://)?liveweb.arte.tv/(?P<lang>fr|de)/(?P<subpage>.+?)/(?P<name>.+)' +    _VIDEOS_URL = r'(?:http://)?videos\.arte\.tv/(?P<lang>fr|de)/.*-(?P<id>.*?)\.html' +    _LIVEWEB_URL = r'(?:http://)?liveweb\.arte\.tv/(?P<lang>fr|de)/(?P<subpage>.+?)/(?P<name>.+)'      _LIVE_URL = r'index-[0-9]+\.html$'      IE_NAME = u'arte.tv' @@ -142,7 +143,9 @@ class ArteTVPlus7IE(InfoExtractor):      def _extract_from_webpage(self, webpage, video_id, lang):          json_url = self._html_search_regex(r'arte_vp_url="(.*?)"', webpage, 'json url') +        return self._extract_from_json_url(json_url, video_id, lang) +    def _extract_from_json_url(self, json_url, video_id, lang):          json_info = self._download_webpage(json_url, video_id, 'Downloading info json')          self.report_extraction(video_id)          info = json.loads(json_info) @@ -257,3 +260,35 @@ class ArteTVFutureIE(ArteTVPlus7IE):          webpage = self._download_webpage(url, anchor_id)          row = get_element_by_id(anchor_id, webpage)          return self._extract_from_webpage(row, anchor_id, lang) + + +class ArteTVDDCIE(ArteTVPlus7IE): +    IE_NAME = u'arte.tv:ddc' +    _VALID_URL = r'http?://ddc\.arte\.tv/(?P<lang>emission|folge)/(?P<id>.+)' + +    _TEST = { +        u'url': u'http://ddc.arte.tv/folge/neues-aus-mauretanien', +        u'file': u'049881-009_PLUS7-D.flv', +        u'info_dict': { +            u'title': u'Mit offenen Karten', +            u'description': u'md5:57929b0eaeddeb8a0c983f58e9ebd3b6', +            u'upload_date': u'20131207', +        }, +        u'params': { +            # rtmp download +            u'skip_download': True, +        }, +    } + +    def _real_extract(self, url): +        video_id, lang = self._extract_url_info(url) +        if lang == 'folge': +            lang = 'de' +        elif lang == 'emission': +            lang = 'fr' +        webpage = self._download_webpage(url, video_id) +        scriptElement = get_element_by_attribute('class', 'visu_video_block', webpage) +        script_url = self._html_search_regex(r'src="(.*?)"', scriptElement, 'script url') +        javascriptPlayerGenerator = self._download_webpage(script_url, video_id, 'Download javascript player generator') +        json_url = self._search_regex(r"json_url=(.*)&rendering_place.*", javascriptPlayerGenerator, 'json url') +        return self._extract_from_json_url(json_url, video_id, lang)  | 
