akkoma/lib/pleroma/web/rich_media/parsers/meta_tags_parser.ex

47 lines
1.3 KiB
Elixir
Raw Normal View History

# Pleroma: A lightweight social networking server
# Copyright © 2017-2020 Pleroma Authors <https://pleroma.social/>
# SPDX-License-Identifier: AGPL-3.0-only
2019-01-10 18:09:56 +00:00
defmodule Pleroma.Web.RichMedia.Parsers.MetaTagsParser do
2020-06-11 13:57:31 +00:00
def parse(data, html, prefix, key_name, value_name \\ "content") do
html
|> get_elements(key_name, prefix)
|> Enum.reduce(data, fn el, acc ->
attributes = normalize_attributes(el, prefix, key_name, value_name)
Map.merge(acc, attributes)
end)
|> maybe_put_title(html)
2019-01-10 18:09:56 +00:00
end
defp get_elements(html, key_name, prefix) do
html |> Floki.find("meta[#{key_name}^='#{prefix}:']")
end
defp normalize_attributes(html_node, prefix, key_name, value_name) do
{_tag, attributes, _children} = html_node
data =
2020-06-09 17:49:24 +00:00
Map.new(attributes, fn {name, value} ->
2019-01-10 18:09:56 +00:00
{name, String.trim_leading(value, "#{prefix}:")}
end)
2020-06-09 17:49:24 +00:00
%{data[key_name] => data[value_name]}
2019-01-10 18:09:56 +00:00
end
2020-06-09 17:49:24 +00:00
defp maybe_put_title(%{"title" => _} = meta, _), do: meta
defp maybe_put_title(meta, html) when meta != %{} do
case get_page_title(html) do
"" -> meta
2020-06-09 17:49:24 +00:00
title -> Map.put_new(meta, "title", title)
end
end
defp maybe_put_title(meta, _), do: meta
defp get_page_title(html) do
2020-01-29 08:13:34 +00:00
Floki.find(html, "html head title") |> List.first() |> Floki.text()
end
2019-01-10 18:09:56 +00:00
end