django/tests/utils_tests/test_html.py

# -*- coding: utf-8 -*-
from __future__ import unicode_literals

import os
from datetime import datetime
from unittest import TestCase

from django.test import ignore_warnings
from django.utils import html, safestring
from django.utils._os import upath
from django.utils.deprecation import RemovedInDjango20Warning
from django.utils.encoding import force_text


class TestUtilsHtml(TestCase):

    def check_output(self, function, value, output=None):
        """
        Check that function(value) equals output.  If output is None,
        check that function(value) equals value.
        """
        if output is None:
            output = value
        self.assertEqual(function(value), output)

    def test_escape(self):
        f = html.escape
        items = (
            ('&', '&amp;'),
            ('<', '&lt;'),
            ('>', '&gt;'),
            ('"', '&quot;'),
            ("'", '&#39;'),
        )
        # Substitution patterns for testing the above items.
        patterns = ("%s", "asdf%sfdsa", "%s1", "1%sb")
        for value, output in items:
            for pattern in patterns:
                self.check_output(f, pattern % value, pattern % output)
            # Check repeated values.
            self.check_output(f, value * 2, output * 2)
        # Verify it doesn't double replace &.
        self.check_output(f, '<&', '&lt;&amp;')

    def test_format_html(self):
        self.assertEqual(
            html.format_html("{} {} {third} {fourth}",
                             "< Dangerous >",
                             html.mark_safe("<b>safe</b>"),
                             third="< dangerous again",
                             fourth=html.mark_safe("<i>safe again</i>")
                             ),
            "&lt; Dangerous &gt; <b>safe</b> &lt; dangerous again <i>safe again</i>"
        )

    def test_linebreaks(self):
        f = html.linebreaks
        items = (
            ("para1\n\npara2\r\rpara3", "<p>para1</p>\n\n<p>para2</p>\n\n<p>para3</p>"),
            ("para1\nsub1\rsub2\n\npara2", "<p>para1<br />sub1<br />sub2</p>\n\n<p>para2</p>"),
            ("para1\r\n\r\npara2\rsub1\r\rpara4", "<p>para1</p>\n\n<p>para2<br />sub1</p>\n\n<p>para4</p>"),
            ("para1\tmore\n\npara2", "<p>para1\tmore</p>\n\n<p>para2</p>"),
        )
        for value, output in items:
            self.check_output(f, value, output)

    def test_strip_tags(self):
        f = html.strip_tags
        items = (
            ('<p>See: &#39;&eacute; is an apostrophe followed by e acute</p>',
             'See: &#39;&eacute; is an apostrophe followed by e acute'),
            ('<adf>a', 'a'),
            ('</adf>a', 'a'),
            ('<asdf><asdf>e', 'e'),
            ('hi, <f x', 'hi, <f x'),
            ('234<235, right?', '234<235, right?'),
            ('a4<a5 right?', 'a4<a5 right?'),
            ('b7>b2!', 'b7>b2!'),
            ('</fe', '</fe'),
            ('<x>b<y>', 'b'),
            ('a<p onclick="alert(\'<test>\')">b</p>c', 'abc'),
            ('a<p a >b</p>c', 'abc'),
            ('d<a:b c:d>e</p>f', 'def'),
            ('<strong>foo</strong><a href="http://example.com">bar</a>', 'foobar'),
            # caused infinite loop on Pythons not patched with
            # http://bugs.python.org/issue20288
            ('&gotcha&#;<>', '&gotcha&#;<>'),
        )
        for value, output in items:
            self.check_output(f, value, output)

        # Some convoluted syntax for which parsing may differ between python versions
        output = html.strip_tags('<sc<!-- -->ript>test<<!-- -->/script>')
        self.assertNotIn('<script>', output)
        self.assertIn('test', output)
        output = html.strip_tags('<script>alert()</script>&h')
        self.assertNotIn('<script>', output)
        self.assertIn('alert()', output)

        # Test with more lengthy content (also catching performance regressions)
        for filename in ('strip_tags1.html', 'strip_tags2.txt'):
            path = os.path.join(os.path.dirname(upath(__file__)), 'files', filename)
            with open(path, 'r') as fp:
                content = force_text(fp.read())
                start = datetime.now()
                stripped = html.strip_tags(content)
                elapsed = datetime.now() - start
            self.assertEqual(elapsed.seconds, 0)
            self.assertIn("Please try again.", stripped)
            self.assertNotIn('<', stripped)

    def test_strip_spaces_between_tags(self):
        f = html.strip_spaces_between_tags
        # Strings that should come out untouched.
        items = (' <adf>', '<adf> ', ' </adf> ', ' <f> x</f>')
        for value in items:
            self.check_output(f, value)
        # Strings that have spaces to strip.
        items = (
            ('<d> </d>', '<d></d>'),
            ('<p>hello </p>\n<p> world</p>', '<p>hello </p><p> world</p>'),
            ('\n<p>\t</p>\n<p> </p>\n', '\n<p></p><p></p>\n'),
        )
        for value, output in items:
            self.check_output(f, value, output)

    @ignore_warnings(category=RemovedInDjango20Warning)
    def test_strip_entities(self):
        f = html.strip_entities
        # Strings that should come out untouched.
        values = ("&", "&a", "&a", "a&#a")
        for value in values:
            self.check_output(f, value)
        # Valid entities that should be stripped from the patterns.
        entities = ("&#1;", "&#12;", "&a;", "&fdasdfasdfasdf;")
        patterns = (
            ("asdf %(entity)s ", "asdf  "),
            ("%(entity)s%(entity)s", ""),
            ("&%(entity)s%(entity)s", "&"),
            ("%(entity)s3", "3"),
        )
        for entity in entities:
            for in_pattern, output in patterns:
                self.check_output(f, in_pattern % {'entity': entity}, output)

    def test_escapejs(self):
        f = html.escapejs
        items = (
            ('"double quotes" and \'single quotes\'', '\\u0022double quotes\\u0022 and \\u0027single quotes\\u0027'),
            (r'\ : backslashes, too', '\\u005C : backslashes, too'),
            ('and lots of whitespace: \r\n\t\v\f\b', 'and lots of whitespace: \\u000D\\u000A\\u0009\\u000B\\u000C\\u0008'),
            (r'<script>and this</script>', '\\u003Cscript\\u003Eand this\\u003C/script\\u003E'),
            ('paragraph separator:\u2029and line separator:\u2028', 'paragraph separator:\\u2029and line separator:\\u2028'),
        )
        for value, output in items:
            self.check_output(f, value, output)

    @ignore_warnings(category=RemovedInDjango20Warning)
    def test_remove_tags(self):
        f = html.remove_tags
        items = (
            ("<b><i>Yes</i></b>", "b i", "Yes"),
            ("<a>x</a> <p><b>y</b></p>", "a b", "x <p>y</p>"),
        )
        for value, tags, output in items:
            self.assertEqual(f(value, tags), output)

    def test_smart_urlquote(self):
        quote = html.smart_urlquote
        # Ensure that IDNs are properly quoted
        self.assertEqual(quote('http://öäü.com/'), 'http://xn--4ca9at.com/')
        self.assertEqual(quote('http://öäü.com/öäü/'), 'http://xn--4ca9at.com/%C3%B6%C3%A4%C3%BC/')
        # Ensure that everything unsafe is quoted, !*'();:@&=+$,/?#[]~ is considered safe as per RFC
        self.assertEqual(quote('http://example.com/path/öäü/'), 'http://example.com/path/%C3%B6%C3%A4%C3%BC/')
        self.assertEqual(quote('http://example.com/%C3%B6/ä/'), 'http://example.com/%C3%B6/%C3%A4/')
        self.assertEqual(quote('http://example.com/?x=1&y=2+3&z='), 'http://example.com/?x=1&y=2+3&z=')
        self.assertEqual(quote('http://example.com/?x=<>"\''), 'http://example.com/?x=%3C%3E%22%27')
        self.assertEqual(quote('http://example.com/?q=http://example.com/?x=1%26q=django'),
                         'http://example.com/?q=http%3A%2F%2Fexample.com%2F%3Fx%3D1%26q%3Ddjango')
        self.assertEqual(quote('http://example.com/?q=http%3A%2F%2Fexample.com%2F%3Fx%3D1%26q%3Ddjango'),
                         'http://example.com/?q=http%3A%2F%2Fexample.com%2F%3Fx%3D1%26q%3Ddjango')

    def test_conditional_escape(self):
        s = '<h1>interop</h1>'
        self.assertEqual(html.conditional_escape(s),
                         '&lt;h1&gt;interop&lt;/h1&gt;')
        self.assertEqual(html.conditional_escape(safestring.mark_safe(s)), s)
Simplified smart_urlquote and added some basic tests. 2013-07-28 08:05:39 +00:00			`# -- coding: utf-8 --`
Fixed #18269 -- Applied unicode_literals for Python 3 compatibility. Thanks Vinay Sajip for the support of his django3 branch and Jannis Leidel for the review. 2012-06-07 16:08:47 +00:00			`from __future__ import unicode_literals`

Added more tests for strip_tags utility Refs #19237. 2013-04-01 14:42:31 +00:00			`import os`
Sorted imports with isort; refs #23860. 2015-01-28 12:35:27 +00:00			`from datetime import datetime`
Stopped using django.utils.unittest in the test suite. Refs #20680. 2013-07-01 12:22:27 +00:00			`from unittest import TestCase`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00
Applied ignore_warnings to Django tests 2014-12-21 20:19:05 +00:00			`from django.test import ignore_warnings`
Fixed #7261 -- support for __html__ for library interoperability The idea is that if an object implements __html__ which returns a string this is used as HTML representation (eg: on escaping). If the object is a str or unicode subclass and returns itself the object is a safe string type. This is an updated patch based on jbalogh and ivank patches. 2013-10-14 22:40:52 +00:00			`from django.utils import html, safestring`
Added more tests for strip_tags utility Refs #19237. 2013-04-01 14:42:31 +00:00			`from django.utils._os import upath`
Applied ignore_warnings to Django tests 2014-12-21 20:19:05 +00:00			`from django.utils.deprecation import RemovedInDjango20Warning`
Fixed #19237 -- Used HTML parser to strip tags The regex method used until now for the strip_tags utility is fast, but subject to flaws and security issues. Consensus and good practice lead use to use a slower but safer method. 2013-05-22 15:29:16 +00:00			`from django.utils.encoding import force_text`
Added more tests for strip_tags utility Refs #19237. 2013-04-01 14:42:31 +00:00
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00
Imported unittest from django.utils in util_tests Without this, the 'new' assertion methods are not present with Python 2.6. 2013-04-01 17:58:16 +00:00			`class TestUtilsHtml(TestCase):`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00
			`def check_output(self, function, value, output=None):`
			`"""`
			`Check that function(value) equals output. If output is None,`
			`check that function(value) equals value.`
			`"""`
			`if output is None:`
			`output = value`
			`self.assertEqual(function(value), output)`

			`def test_escape(self):`
			`f = html.escape`
			`items = (`
Fix all violators of E231 2013-10-26 19:15:03 +00:00			`('&', '&'),`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00			`('<', '<'),`
			`('>', '>'),`
			`('"', '"'),`
			`("'", '''),`
			`)`
			`# Substitution patterns for testing the above items.`
			`patterns = ("%s", "asdf%sfdsa", "%s1", "1%sb")`
			`for value, output in items:`
			`for pattern in patterns:`
			`self.check_output(f, pattern % value, pattern % output)`
			`# Check repeated values.`
			`self.check_output(f, value * 2, output * 2)`
			`# Verify it doesn't double replace &.`
			`self.check_output(f, '<&', '<&')`

Added 'format_html' utility for formatting HTML fragments safely 2012-06-30 17:54:38 +00:00			`def test_format_html(self):`
			`self.assertEqual(`
Removed redundant numbered parameters from str.format(). Since Python 2.7 and 3.1, "{0} {1}" is equivalent to "{} {}". 2014-11-27 00:41:27 +00:00			`html.format_html("{} {} {third} {fourth}",`
Removed u prefixes on unicode strings. They break Python 3. 2012-07-20 10:29:22 +00:00			`"< Dangerous >",`
			`html.mark_safe("<b>safe</b>"),`
Added 'format_html' utility for formatting HTML fragments safely 2012-06-30 17:54:38 +00:00			`third="< dangerous again",`
Removed u prefixes on unicode strings. They break Python 3. 2012-07-20 10:29:22 +00:00			`fourth=html.mark_safe("<i>safe again</i>")`
Added 'format_html' utility for formatting HTML fragments safely 2012-06-30 17:54:38 +00:00			`),`
Removed u prefixes on unicode strings. They break Python 3. 2012-07-20 10:29:22 +00:00			`"< Dangerous > <b>safe</b> < dangerous again <i>safe again</i>"`
Fixed #21287 -- Fixed E123 pep8 warnings 2013-10-18 09:02:43 +00:00			`)`
Added 'format_html' utility for formatting HTML fragments safely 2012-06-30 17:54:38 +00:00
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00			`def test_linebreaks(self):`
			`f = html.linebreaks`
			`items = (`
			`("para1\n\npara2\r\rpara3", "<p>para1</p>\n\n<p>para2</p>\n\n<p>para3</p>"),`
			`("para1\nsub1\rsub2\n\npara2", "<p>para1<br />sub1<br />sub2</p>\n\n<p>para2</p>"),`
			`("para1\r\n\r\npara2\rsub1\r\rpara4", "<p>para1</p>\n\n<p>para2<br />sub1</p>\n\n<p>para4</p>"),`
			`("para1\tmore\n\npara2", "<p>para1\tmore</p>\n\n<p>para2</p>"),`
			`)`
			`for value, output in items:`
			`self.check_output(f, value, output)`

			`def test_strip_tags(self):`
			`f = html.strip_tags`
			`items = (`
Fixed #19237 -- Used HTML parser to strip tags The regex method used until now for the strip_tags utility is fast, but subject to flaws and security issues. Consensus and good practice lead use to use a slower but safer method. 2013-05-22 15:29:16 +00:00			`('<p>See: 'é is an apostrophe followed by e acute</p>',`
			`'See: 'é is an apostrophe followed by e acute'),`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00			`('<adf>a', 'a'),`
			`('</adf>a', 'a'),`
			`('<asdf><asdf>e', 'e'),`
Fixed #19237 -- Used HTML parser to strip tags The regex method used until now for the strip_tags utility is fast, but subject to flaws and security issues. Consensus and good practice lead use to use a slower but safer method. 2013-05-22 15:29:16 +00:00			`('hi, <f x', 'hi, <f x'),`
Fixed #19237 (again) - Made strip_tags consistent between Python versions 2013-05-23 12:00:17 +00:00			`('234<235, right?', '234<235, right?'),`
			`('a4<a5 right?', 'a4<a5 right?'),`
			`('b7>b2!', 'b7>b2!'),`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00			`('</fe', '</fe'),`
			`('<x>b<y>', 'b'),`
Fixed #19237 -- Improved strip_tags utility The previous pattern didn't properly addressed cases where '>' was present inside quoted tag content. 2012-11-24 11:10:25 +00:00			`('a<p onclick="alert(\'<test>\')">b</p>c', 'abc'),`
			`('a<p a >b</p>c', 'abc'),`
			`('d<a:b c:d>e</p>f', 'def'),`
Improved regex in strip_tags Thanks Pablo Recio for the report. Refs #19237. 2013-02-06 20:20:43 +00:00			`('<strong>foo</strong><a href="http://example.com">bar</a>', 'foobar'),`
Fixed an infinite loop possibility in strip_tags(). This is a security fix; disclosure to follow shortly. 2015-03-04 13:11:25 +00:00			`# caused infinite loop on Pythons not patched with`
			`# http://bugs.python.org/issue20288`
			`('&gotcha&#;<>', '&gotcha&#;<>'),`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00			`)`
			`for value, output in items:`
			`self.check_output(f, value, output)`

Tweaked strip_tags tests to pass on Python 3.3 2014-03-22 13:41:45 +00:00			`# Some convoluted syntax for which parsing may differ between python versions`
			`output = html.strip_tags('<sc<!-- -->ript>test<<!-- -->/script>')`
			`self.assertNotIn('<script>', output)`
			`self.assertIn('test', output)`
			`output = html.strip_tags('<script>alert()</script>&h')`
			`self.assertNotIn('<script>', output)`
			`self.assertIn('alert()', output)`

Added more tests for strip_tags utility Refs #19237. 2013-04-01 14:42:31 +00:00			`# Test with more lengthy content (also catching performance regressions)`
			`for filename in ('strip_tags1.html', 'strip_tags2.txt'):`
			`path = os.path.join(os.path.dirname(upath(__file__)), 'files', filename)`
			`with open(path, 'r') as fp:`
Fixed #19237 -- Used HTML parser to strip tags The regex method used until now for the strip_tags utility is fast, but subject to flaws and security issues. Consensus and good practice lead use to use a slower but safer method. 2013-05-22 15:29:16 +00:00			`content = force_text(fp.read())`
Added more tests for strip_tags utility Refs #19237. 2013-04-01 14:42:31 +00:00			`start = datetime.now()`
Fixed #19237 -- Used HTML parser to strip tags The regex method used until now for the strip_tags utility is fast, but subject to flaws and security issues. Consensus and good practice lead use to use a slower but safer method. 2013-05-22 15:29:16 +00:00			`stripped = html.strip_tags(content)`
Added more tests for strip_tags utility Refs #19237. 2013-04-01 14:42:31 +00:00			`elapsed = datetime.now() - start`
			`self.assertEqual(elapsed.seconds, 0)`
			`self.assertIn("Please try again.", stripped)`
			`self.assertNotIn('<', stripped)`

Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00			`def test_strip_spaces_between_tags(self):`
			`f = html.strip_spaces_between_tags`
			`# Strings that should come out untouched.`
			`items = (' <adf>', '<adf> ', ' </adf> ', ' <f> x</f>')`
			`for value in items:`
			`self.check_output(f, value)`
			`# Strings that have spaces to strip.`
			`items = (`
			`('<d> </d>', '<d></d>'),`
			`('<p>hello </p>\n<p> world</p>', '<p>hello </p><p> world</p>'),`
			`('\n<p>\t</p>\n<p> </p>\n', '\n<p></p><p></p>\n'),`
			`)`
			`for value, output in items:`
			`self.check_output(f, value, output)`

Applied ignore_warnings to Django tests 2014-12-21 20:19:05 +00:00			`@ignore_warnings(category=RemovedInDjango20Warning)`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00			`def test_strip_entities(self):`
			`f = html.strip_entities`
			`# Strings that should come out untouched.`
			`values = ("&", "&a", "&a", "a&#a")`
			`for value in values:`
Applied ignore_warnings to Django tests 2014-12-21 20:19:05 +00:00			`self.check_output(f, value)`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00			`# Valid entities that should be stripped from the patterns.`
			`entities = ("", "", "&a;", "&fdasdfasdfasdf;")`
			`patterns = (`
			`("asdf %(entity)s ", "asdf "),`
			`("%(entity)s%(entity)s", ""),`
			`("&%(entity)s%(entity)s", "&"),`
			`("%(entity)s3", "3"),`
			`)`
			`for entity in entities:`
			`for in_pattern, output in patterns:`
Applied ignore_warnings to Django tests 2014-12-21 20:19:05 +00:00			`self.check_output(f, in_pattern % {'entity': entity}, output)`
Reorganized utils tests so it's all in separate modules. Thanks to Stephan Jaekel. git-svn-id: http://code.djangoproject.com/svn/django/trunk@13889 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2010-09-27 15:15:04 +00:00
Fixed #2986 -- Made the JavaScript code that drives related model instance addition in a popup window handle a model representation containing new lines. Also, moved the escapejs functionality yoo django.utils.html so it can be used from Python code. Thanks andrewwatts for the patch. git-svn-id: http://code.djangoproject.com/svn/django/trunk@15131 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2011-01-02 17:34:52 +00:00			`def test_escapejs(self):`
			`f = html.escapejs`
			`items = (`
Fixed #18269 -- Applied unicode_literals for Python 3 compatibility. Thanks Vinay Sajip for the support of his django3 branch and Jannis Leidel for the review. 2012-06-07 16:08:47 +00:00			`('"double quotes" and \'single quotes\'', '\\u0022double quotes\\u0022 and \\u0027single quotes\\u0027'),`
			`(r'\ : backslashes, too', '\\u005C : backslashes, too'),`
			`('and lots of whitespace: \r\n\t\v\f\b', 'and lots of whitespace: \\u000D\\u000A\\u0009\\u000B\\u000C\\u0008'),`
			`(r'<script>and this</script>', '\\u003Cscript\\u003Eand this\\u003C/script\\u003E'),`
			`('paragraph separator:\u2029and line separator:\u2028', 'paragraph separator:\\u2029and line separator:\\u2028'),`
Fixed #2986 -- Made the JavaScript code that drives related model instance addition in a popup window handle a model representation containing new lines. Also, moved the escapejs functionality yoo django.utils.html so it can be used from Python code. Thanks andrewwatts for the patch. git-svn-id: http://code.djangoproject.com/svn/django/trunk@15131 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2011-01-02 17:34:52 +00:00			`)`
			`for value, output in items:`
			`self.check_output(f, value, output)`
Fixed #7267 - UnicodeDecodeError in clean_html Thanks to Nikolay for the report, and gav and aaugustin for the patch. git-svn-id: http://code.djangoproject.com/svn/django/trunk@16118 bcc190cf-cafb-0310-a4f2-bffc1f526a37 2011-04-28 14:08:53 +00:00
Applied ignore_warnings to Django tests 2014-12-21 20:19:05 +00:00			`@ignore_warnings(category=RemovedInDjango20Warning)`
Fixed #14516 -- Extract methods from removetags and slugify template filters Patch by @jphalip updated to apply, documentation and release notes added. I've documented strip_tags as well as remove_tags as the difference between the two wouldn't be immediately obvious. 2012-08-18 12:53:22 +00:00			`def test_remove_tags(self):`
			`f = html.remove_tags`
			`items = (`
			`("<b><i>Yes</i></b>", "b i", "Yes"),`
			`("<a>x</a> <p><b>y</b></p>", "a b", "x <p>y</p>"),`
			`)`
			`for value, tags, output in items:`
Applied ignore_warnings to Django tests 2014-12-21 20:19:05 +00:00			`self.assertEqual(f(value, tags), output)`
Simplified smart_urlquote and added some basic tests. 2013-07-28 08:05:39 +00:00
			`def test_smart_urlquote(self):`
			`quote = html.smart_urlquote`
			`# Ensure that IDNs are properly quoted`
			`self.assertEqual(quote('http://öäü.com/'), 'http://xn--4ca9at.com/')`
			`self.assertEqual(quote('http://öäü.com/öäü/'), 'http://xn--4ca9at.com/%C3%B6%C3%A4%C3%BC/')`
			`# Ensure that everything unsafe is quoted, !*'();:@&=+$,/?#[]~ is considered safe as per RFC`
			`self.assertEqual(quote('http://example.com/path/öäü/'), 'http://example.com/path/%C3%B6%C3%A4%C3%BC/')`
			`self.assertEqual(quote('http://example.com/%C3%B6/ä/'), 'http://example.com/%C3%B6/%C3%A4/')`
Fixed #22267 -- Fixed unquote/quote in smart_urlquote Thanks Md. Enzam Hossain for the report and initial patch, and Tim Graham for the review. 2014-06-26 19:14:30 +00:00			`self.assertEqual(quote('http://example.com/?x=1&y=2+3&z='), 'http://example.com/?x=1&y=2+3&z=')`
Fixed urlize after smart_urlquote rewrite Refs #22267. 2014-08-09 10:44:48 +00:00			`self.assertEqual(quote('http://example.com/?x=<>"\''), 'http://example.com/?x=%3C%3E%22%27')`
Fixed #22267 -- Fixed unquote/quote in smart_urlquote Thanks Md. Enzam Hossain for the report and initial patch, and Tim Graham for the review. 2014-06-26 19:14:30 +00:00			`self.assertEqual(quote('http://example.com/?q=http://example.com/?x=1%26q=django'),`
			`'http://example.com/?q=http%3A%2F%2Fexample.com%2F%3Fx%3D1%26q%3Ddjango')`
			`self.assertEqual(quote('http://example.com/?q=http%3A%2F%2Fexample.com%2F%3Fx%3D1%26q%3Ddjango'),`
			`'http://example.com/?q=http%3A%2F%2Fexample.com%2F%3Fx%3D1%26q%3Ddjango')`
Fixed #7261 -- support for __html__ for library interoperability The idea is that if an object implements __html__ which returns a string this is used as HTML representation (eg: on escaping). If the object is a str or unicode subclass and returns itself the object is a safe string type. This is an updated patch based on jbalogh and ivank patches. 2013-10-14 22:40:52 +00:00
			`def test_conditional_escape(self):`
			`s = '<h1>interop</h1>'`
			`self.assertEqual(html.conditional_escape(s),`
			`'<h1>interop</h1>')`
			`self.assertEqual(html.conditional_escape(safestring.mark_safe(s)), s)`