from unstructured.partition.html import partition_html def test_alternative_image_text_can_be_included(): # language=HTML html = """
ALT TEXT Logo
""" _, image_to_text_alt_mode = partition_html( text=html, image_alt_mode="to_text", html_parser_version="v2", ) assert "ALT TEXT Logo" in image_to_text_alt_mode.text _, image_none_alt_mode = partition_html( text=html, image_alt_mode=None, html_parser_version="v2", ) assert "ALT TEXT Logo" not in image_none_alt_mode.text def test_alternative_image_text_can_be_included_when_nested_in_paragraph(): # language=HTML html = """

ALT TEXT Logo

""" _, paragraph_to_text_alt_mode = partition_html( text=html, image_alt_mode="to_text", html_parser_version="v2", ) assert "ALT TEXT Logo" in paragraph_to_text_alt_mode.text _, paragraph_none_alt_mode = partition_html( text=html, image_alt_mode=None, html_parser_version="v2", ) assert "ALT TEXT Logo" not in paragraph_none_alt_mode.text