diff options
Diffstat (limited to 'test')
| -rw-r--r-- | test/helper.py | 129 | ||||
| -rw-r--r-- | test/test_InfoExtractor.py | 8 | ||||
| -rw-r--r-- | test/test_compat.py | 17 | ||||
| -rw-r--r-- | test/test_download.py | 2 | ||||
| -rw-r--r-- | test/test_jsinterp.py | 3 | ||||
| -rw-r--r-- | test/test_subtitles.py | 21 | ||||
| -rw-r--r-- | test/test_utils.py | 24 | ||||
| -rw-r--r-- | test/test_youtube_lists.py | 9 | 
8 files changed, 151 insertions, 62 deletions
| diff --git a/test/helper.py b/test/helper.py index cb6eec8d9..bdd7acca4 100644 --- a/test/helper.py +++ b/test/helper.py @@ -89,66 +89,81 @@ def gettestcases(include_onlymatching=False):  md5 = lambda s: hashlib.md5(s.encode('utf-8')).hexdigest() -def expect_info_dict(self, got_dict, expected_dict): +def expect_value(self, got, expected, field): +    if isinstance(expected, compat_str) and expected.startswith('re:'): +        match_str = expected[len('re:'):] +        match_rex = re.compile(match_str) + +        self.assertTrue( +            isinstance(got, compat_str), +            'Expected a %s object, but got %s for field %s' % ( +                compat_str.__name__, type(got).__name__, field)) +        self.assertTrue( +            match_rex.match(got), +            'field %s (value: %r) should match %r' % (field, got, match_str)) +    elif isinstance(expected, compat_str) and expected.startswith('startswith:'): +        start_str = expected[len('startswith:'):] +        self.assertTrue( +            isinstance(got, compat_str), +            'Expected a %s object, but got %s for field %s' % ( +                compat_str.__name__, type(got).__name__, field)) +        self.assertTrue( +            got.startswith(start_str), +            'field %s (value: %r) should start with %r' % (field, got, start_str)) +    elif isinstance(expected, compat_str) and expected.startswith('contains:'): +        contains_str = expected[len('contains:'):] +        self.assertTrue( +            isinstance(got, compat_str), +            'Expected a %s object, but got %s for field %s' % ( +                compat_str.__name__, type(got).__name__, field)) +        self.assertTrue( +            contains_str in got, +            'field %s (value: %r) should contain %r' % (field, got, contains_str)) +    elif isinstance(expected, type): +        self.assertTrue( +            isinstance(got, expected), +            'Expected type %r for field %s, but got value %r of type %r' % (expected, field, got, type(got))) +    elif isinstance(expected, dict) and isinstance(got, dict): +        expect_dict(self, got, expected) +    elif isinstance(expected, list) and isinstance(got, list): +        self.assertEqual( +            len(expected), len(got), +            'Expect a list of length %d, but got a list of length %d for field %s' % ( +                len(expected), len(got), field)) +        for index, (item_got, item_expected) in enumerate(zip(got, expected)): +            type_got = type(item_got) +            type_expected = type(item_expected) +            self.assertEqual( +                type_expected, type_got, +                'Type mismatch for list item at index %d for field %s, expected %r, got %r' % ( +                    index, field, type_expected, type_got)) +            expect_value(self, item_got, item_expected, field) +    else: +        if isinstance(expected, compat_str) and expected.startswith('md5:'): +            got = 'md5:' + md5(got) +        elif isinstance(expected, compat_str) and expected.startswith('mincount:'): +            self.assertTrue( +                isinstance(got, (list, dict)), +                'Expected field %s to be a list or a dict, but it is of type %s' % ( +                    field, type(got).__name__)) +            expected_num = int(expected.partition(':')[2]) +            assertGreaterEqual( +                self, len(got), expected_num, +                'Expected %d items in field %s, but only got %d' % (expected_num, field, len(got))) +            return +        self.assertEqual( +            expected, got, +            'Invalid value for field %s, expected %r, got %r' % (field, expected, got)) + + +def expect_dict(self, got_dict, expected_dict):      for info_field, expected in expected_dict.items(): -        if isinstance(expected, compat_str) and expected.startswith('re:'): -            got = got_dict.get(info_field) -            match_str = expected[len('re:'):] -            match_rex = re.compile(match_str) +        got = got_dict.get(info_field) +        expect_value(self, got, expected, info_field) -            self.assertTrue( -                isinstance(got, compat_str), -                'Expected a %s object, but got %s for field %s' % ( -                    compat_str.__name__, type(got).__name__, info_field)) -            self.assertTrue( -                match_rex.match(got), -                'field %s (value: %r) should match %r' % (info_field, got, match_str)) -        elif isinstance(expected, compat_str) and expected.startswith('startswith:'): -            got = got_dict.get(info_field) -            start_str = expected[len('startswith:'):] -            self.assertTrue( -                isinstance(got, compat_str), -                'Expected a %s object, but got %s for field %s' % ( -                    compat_str.__name__, type(got).__name__, info_field)) -            self.assertTrue( -                got.startswith(start_str), -                'field %s (value: %r) should start with %r' % (info_field, got, start_str)) -        elif isinstance(expected, compat_str) and expected.startswith('contains:'): -            got = got_dict.get(info_field) -            contains_str = expected[len('contains:'):] -            self.assertTrue( -                isinstance(got, compat_str), -                'Expected a %s object, but got %s for field %s' % ( -                    compat_str.__name__, type(got).__name__, info_field)) -            self.assertTrue( -                contains_str in got, -                'field %s (value: %r) should contain %r' % (info_field, got, contains_str)) -        elif isinstance(expected, type): -            got = got_dict.get(info_field) -            self.assertTrue(isinstance(got, expected), -                            'Expected type %r for field %s, but got value %r of type %r' % (expected, info_field, got, type(got))) -        else: -            if isinstance(expected, compat_str) and expected.startswith('md5:'): -                got = 'md5:' + md5(got_dict.get(info_field)) -            elif isinstance(expected, compat_str) and expected.startswith('mincount:'): -                got = got_dict.get(info_field) -                self.assertTrue( -                    isinstance(got, (list, dict)), -                    'Expected field %s to be a list or a dict, but it is of type %s' % ( -                        info_field, type(got).__name__)) -                expected_num = int(expected.partition(':')[2]) -                assertGreaterEqual( -                    self, len(got), expected_num, -                    'Expected %d items in field %s, but only got %d' % ( -                        expected_num, info_field, len(got) -                    ) -                ) -                continue -            else: -                got = got_dict.get(info_field) -            self.assertEqual(expected, got, -                             'invalid value for field %s, expected %r, got %r' % (info_field, expected, got)) +def expect_info_dict(self, got_dict, expected_dict): +    expect_dict(self, got_dict, expected_dict)      # Check for the presence of mandatory fields      if got_dict.get('_type') not in ('playlist', 'multi_video'):          for key in ('id', 'url', 'title', 'ext'): diff --git a/test/test_InfoExtractor.py b/test/test_InfoExtractor.py index be8d12997..938466a80 100644 --- a/test/test_InfoExtractor.py +++ b/test/test_InfoExtractor.py @@ -35,10 +35,18 @@ class TestInfoExtractor(unittest.TestCase):              <meta name="og:title" content='Foo'/>              <meta content="Some video's description " name="og:description"/>              <meta property='og:image' content='http://domain.com/pic.jpg?key1=val1&key2=val2'/> +            <meta content='application/x-shockwave-flash' property='og:video:type'> +            <meta content='Foo' property=og:foobar> +            <meta name="og:test1" content='foo > < bar'/> +            <meta name="og:test2" content="foo >//< bar"/>              '''          self.assertEqual(ie._og_search_title(html), 'Foo')          self.assertEqual(ie._og_search_description(html), 'Some video\'s description ')          self.assertEqual(ie._og_search_thumbnail(html), 'http://domain.com/pic.jpg?key1=val1&key2=val2') +        self.assertEqual(ie._og_search_video_url(html, default=None), None) +        self.assertEqual(ie._og_search_property('foobar', html), 'Foo') +        self.assertEqual(ie._og_search_property('test1', html), 'foo > < bar') +        self.assertEqual(ie._og_search_property('test2', html), 'foo >//< bar')      def test_html_search_meta(self):          ie = self.ie diff --git a/test/test_compat.py b/test/test_compat.py index 4ee0dc99d..b6bfad05e 100644 --- a/test/test_compat.py +++ b/test/test_compat.py @@ -13,8 +13,10 @@ sys.path.insert(0, os.path.dirname(os.path.dirname(os.path.abspath(__file__))))  from youtube_dl.utils import get_filesystem_encoding  from youtube_dl.compat import (      compat_getenv, +    compat_etree_fromstring,      compat_expanduser,      compat_shlex_split, +    compat_str,      compat_urllib_parse_unquote,      compat_urllib_parse_unquote_plus,  ) @@ -71,5 +73,20 @@ class TestCompat(unittest.TestCase):      def test_compat_shlex_split(self):          self.assertEqual(compat_shlex_split('-option "one two"'), ['-option', 'one two']) +    def test_compat_etree_fromstring(self): +        xml = ''' +            <root foo="bar" spam="中文"> +                <normal>foo</normal> +                <chinese>中文</chinese> +                <foo><bar>spam</bar></foo> +            </root> +        ''' +        doc = compat_etree_fromstring(xml.encode('utf-8')) +        self.assertTrue(isinstance(doc.attrib['foo'], compat_str)) +        self.assertTrue(isinstance(doc.attrib['spam'], compat_str)) +        self.assertTrue(isinstance(doc.find('normal').text, compat_str)) +        self.assertTrue(isinstance(doc.find('chinese').text, compat_str)) +        self.assertTrue(isinstance(doc.find('foo/bar').text, compat_str)) +  if __name__ == '__main__':      unittest.main() diff --git a/test/test_download.py b/test/test_download.py index 284418834..a3f1c0644 100644 --- a/test/test_download.py +++ b/test/test_download.py @@ -102,7 +102,7 @@ def generator(test_case):          params = get_params(test_case.get('params', {}))          if is_playlist and 'playlist' not in test_case: -            params.setdefault('extract_flat', True) +            params.setdefault('extract_flat', 'in_playlist')              params.setdefault('skip_download', True)          ydl = YoutubeDL(params, auto_init=False) diff --git a/test/test_jsinterp.py b/test/test_jsinterp.py index fc73e5dc2..63c350b8f 100644 --- a/test/test_jsinterp.py +++ b/test/test_jsinterp.py @@ -19,6 +19,9 @@ class TestJSInterpreter(unittest.TestCase):          jsi = JSInterpreter('function x3(){return 42;}')          self.assertEqual(jsi.call_function('x3'), 42) +        jsi = JSInterpreter('var x5 = function(){return 42;}') +        self.assertEqual(jsi.call_function('x5'), 42) +      def test_calc(self):          jsi = JSInterpreter('function x4(a){return 2*a+1;}')          self.assertEqual(jsi.call_function('x4', 3), 7) diff --git a/test/test_subtitles.py b/test/test_subtitles.py index 0343967d9..75f0ea75f 100644 --- a/test/test_subtitles.py +++ b/test/test_subtitles.py @@ -28,6 +28,7 @@ from youtube_dl.extractor import (      ThePlatformFeedIE,      RTVEALaCartaIE,      FunnyOrDieIE, +    DemocracynowIE,  ) @@ -346,5 +347,25 @@ class TestFunnyOrDieSubtitles(BaseTestSubtitles):          self.assertEqual(md5(subtitles['en']), 'c5593c193eacd353596c11c2d4f9ecc4') +class TestDemocracynowSubtitles(BaseTestSubtitles): +    url = 'http://www.democracynow.org/shows/2015/7/3' +    IE = DemocracynowIE + +    def test_allsubtitles(self): +        self.DL.params['writesubtitles'] = True +        self.DL.params['allsubtitles'] = True +        subtitles = self.getSubtitles() +        self.assertEqual(set(subtitles.keys()), set(['en'])) +        self.assertEqual(md5(subtitles['en']), 'acaca989e24a9e45a6719c9b3d60815c') + +    def test_subtitles_in_page(self): +        self.url = 'http://www.democracynow.org/2015/7/3/this_flag_comes_down_today_bree' +        self.DL.params['writesubtitles'] = True +        self.DL.params['allsubtitles'] = True +        subtitles = self.getSubtitles() +        self.assertEqual(set(subtitles.keys()), set(['en'])) +        self.assertEqual(md5(subtitles['en']), 'acaca989e24a9e45a6719c9b3d60815c') + +  if __name__ == '__main__':      unittest.main() diff --git a/test/test_utils.py b/test/test_utils.py index a5f164c49..01829f71e 100644 --- a/test/test_utils.py +++ b/test/test_utils.py @@ -68,6 +68,9 @@ from youtube_dl.utils import (      cli_valueless_option,      cli_bool_option,  ) +from youtube_dl.compat import ( +    compat_etree_fromstring, +)  class TestUtil(unittest.TestCase): @@ -233,6 +236,7 @@ class TestUtil(unittest.TestCase):              unified_strdate('2/2/2015 6:47:40 PM', day_first=False),              '20150202')          self.assertEqual(unified_strdate('25-09-2014'), '20140925') +        self.assertEqual(unified_strdate('UNKNOWN DATE FORMAT'), None)      def test_find_xpath_attr(self):          testxml = '''<root> @@ -242,7 +246,7 @@ class TestUtil(unittest.TestCase):              <node x="b" y="d" />              <node x="" />          </root>''' -        doc = xml.etree.ElementTree.fromstring(testxml) +        doc = compat_etree_fromstring(testxml)          self.assertEqual(find_xpath_attr(doc, './/fourohfour', 'n'), None)          self.assertEqual(find_xpath_attr(doc, './/fourohfour', 'n', 'v'), None) @@ -263,7 +267,7 @@ class TestUtil(unittest.TestCase):                  <url>http://server.com/download.mp3</url>              </media:song>          </root>''' -        doc = xml.etree.ElementTree.fromstring(testxml) +        doc = compat_etree_fromstring(testxml)          find = lambda p: doc.find(xpath_with_ns(p, {'media': 'http://example.com/'}))          self.assertTrue(find('media:song') is not None)          self.assertEqual(find('media:song/media:author').text, 'The Author') @@ -275,9 +279,16 @@ class TestUtil(unittest.TestCase):          p = xml.etree.ElementTree.SubElement(div, 'p')          p.text = 'Foo'          self.assertEqual(xpath_element(doc, 'div/p'), p) +        self.assertEqual(xpath_element(doc, ['div/p']), p) +        self.assertEqual(xpath_element(doc, ['div/bar', 'div/p']), p)          self.assertEqual(xpath_element(doc, 'div/bar', default='default'), 'default') +        self.assertEqual(xpath_element(doc, ['div/bar'], default='default'), 'default')          self.assertTrue(xpath_element(doc, 'div/bar') is None) +        self.assertTrue(xpath_element(doc, ['div/bar']) is None) +        self.assertTrue(xpath_element(doc, ['div/bar'], 'div/baz') is None)          self.assertRaises(ExtractorError, xpath_element, doc, 'div/bar', fatal=True) +        self.assertRaises(ExtractorError, xpath_element, doc, ['div/bar'], fatal=True) +        self.assertRaises(ExtractorError, xpath_element, doc, ['div/bar', 'div/baz'], fatal=True)      def test_xpath_text(self):          testxml = '''<root> @@ -285,7 +296,7 @@ class TestUtil(unittest.TestCase):                  <p>Foo</p>              </div>          </root>''' -        doc = xml.etree.ElementTree.fromstring(testxml) +        doc = compat_etree_fromstring(testxml)          self.assertEqual(xpath_text(doc, 'div/p'), 'Foo')          self.assertEqual(xpath_text(doc, 'div/bar', default='default'), 'default')          self.assertTrue(xpath_text(doc, 'div/bar') is None) @@ -297,7 +308,7 @@ class TestUtil(unittest.TestCase):                  <p x="a">Foo</p>              </div>          </root>''' -        doc = xml.etree.ElementTree.fromstring(testxml) +        doc = compat_etree_fromstring(testxml)          self.assertEqual(xpath_attr(doc, 'div/p', 'x'), 'a')          self.assertEqual(xpath_attr(doc, 'div/bar', 'x'), None)          self.assertEqual(xpath_attr(doc, 'div/p', 'y'), None) @@ -425,6 +436,8 @@ class TestUtil(unittest.TestCase):          self.assertEqual(parse_iso8601('2014-03-23T22:04:26+0000'), 1395612266)          self.assertEqual(parse_iso8601('2014-03-23T22:04:26Z'), 1395612266)          self.assertEqual(parse_iso8601('2014-03-23T22:04:26.1234Z'), 1395612266) +        self.assertEqual(parse_iso8601('2015-09-29T08:27:31.727'), 1443515251) +        self.assertEqual(parse_iso8601('2015-09-29T08-27-31.727'), None)      def test_strip_jsonp(self):          stripped = strip_jsonp('cb ([ {"id":"532cb",\n\n\n"x":\n3}\n]\n);') @@ -495,6 +508,9 @@ class TestUtil(unittest.TestCase):              "playlist":[{"controls":{"all":null}}]          }''') +        inp = '''"The CW\\'s \\'Crazy Ex-Girlfriend\\'"''' +        self.assertEqual(js_to_json(inp), '''"The CW's 'Crazy Ex-Girlfriend'"''') +          inp = '"SAND Number: SAND 2013-7800P\\nPresenter: Tom Russo\\nHabanero Software Training - Xyce Software\\nXyce, Sandia\\u0027s"'          json_code = js_to_json(inp)          self.assertEqual(json.loads(json_code), json.loads(inp)) diff --git a/test/test_youtube_lists.py b/test/test_youtube_lists.py index c889b6f15..26aadb34f 100644 --- a/test/test_youtube_lists.py +++ b/test/test_youtube_lists.py @@ -57,5 +57,14 @@ class TestYoutubeLists(unittest.TestCase):          entries = result['entries']          self.assertEqual(len(entries), 100) +    def test_youtube_flat_playlist_titles(self): +        dl = FakeYDL() +        dl.params['extract_flat'] = True +        ie = YoutubePlaylistIE(dl) +        result = ie.extract('https://www.youtube.com/playlist?list=PLwiyx1dc3P2JR9N8gQaQN_BCvlSlap7re') +        self.assertIsPlaylist(result) +        for entry in result['entries']: +            self.assertTrue(entry.get('title')) +  if __name__ == '__main__':      unittest.main() | 
