diff options
| author | Sergey M․ <dstftw@gmail.com> | 2015-03-22 07:44:28 +0600 | 
|---|---|---|
| committer | Sergey M․ <dstftw@gmail.com> | 2015-03-22 07:44:28 +0600 | 
| commit | ef249a2cd7a7cfbd92a030cb72e238ba4ad52604 (patch) | |
| tree | b64049a9ec6f311a458211016d6d0f099a792e1a /youtube_dl/extractor/comedycentral.py | |
| parent | a09141548aa31db7c7d9457b10f5c84e6e32beba (diff) | |
| parent | 9ef4f12b534578ae3d3e47815492c90826c03c36 (diff) | |
Merge branch 'the-daily-show-podcast' of https://github.com/fstirlitz/youtube-dl into fstirlitz-the-daily-show-podcast
Diffstat (limited to 'youtube_dl/extractor/comedycentral.py')
| -rw-r--r-- | youtube_dl/extractor/comedycentral.py | 25 | 
1 files changed, 25 insertions, 0 deletions
diff --git a/youtube_dl/extractor/comedycentral.py b/youtube_dl/extractor/comedycentral.py index e5edcc84b..bd3817b56 100644 --- a/youtube_dl/extractor/comedycentral.py +++ b/youtube_dl/extractor/comedycentral.py @@ -2,6 +2,7 @@ from __future__ import unicode_literals  import re +from .common import InfoExtractor  from .mtv import MTVServicesInfoExtractor  from ..compat import (      compat_str, @@ -272,3 +273,27 @@ class ComedyCentralShowsIE(MTVServicesInfoExtractor):              'title': show_name + ' ' + title,              'description': description,          } + +class TheDailyShowPodcastIE(InfoExtractor): +    _VALID_URL = r'(?P<scheme>https?:)?//thedailyshow\.cc\.com/podcast/(?P<id>[a-z\-]+)' +    _TESTS = [{ +        "url": "http://thedailyshow.cc.com/podcast/episodetwelve", +        'only_matching': True, +    }] + +    def _real_extract(self, url): +        display_id = self._match_id(url) +        webpage = self._download_webpage(url, display_id) + +        player_url = self._search_regex(r'<iframe(?:\s+[^>]+)?\s*src="((?:https?:)?//html5-player\.libsyn\.com/embed/episode/id/[0-9]+)', webpage, 'player URL') +        if player_url.startswith('//'): +            mobj = re.match(self._VALID_URL, url) +            scheme = mobj.group('scheme') +            if not scheme: +                scheme = 'https:' +            player_url = scheme + player_url + +        return { +            '_type': 'url_transparent', +            'url': player_url, +        }  | 
