打开/关闭菜单
60
73
29
1.1K
武外梗百科
打开/关闭外观设置菜单
打开/关闭个人菜单
未登录
未登录用户的IP地址会在进行任意编辑后公开展示。

Module:PageSummary:修订间差异

武外梗百科 爱国好学自强图新的百科全书
无编辑摘要
无编辑摘要
 
(未显示18个用户的47个中间版本)
第1行: 第1行:
local p = {}
local p = {}
-- 強化版清理:徹底根除內文圖片,保留加粗和連結
local function cleanContent(text)
    if not text then return "" end
    -- 1. 基礎清理(注釋、腳注、表格、模板)
    text = mw.ustring.gsub(text, "<!%-%-.-%-%->", "")
    text = mw.ustring.gsub(text, "<ref[^>]*>.-</ref>", "")
    text = mw.ustring.gsub(text, "<ref[^>]-/>", "")
   
    -- 移除表格 {| ... |}
    local prev
    repeat
        prev = text
        text = mw.ustring.gsub(text, "{|[^{}]*|}", "")
    until text == prev
    -- 移除模板 {{ ... }}
    local count = 0
    repeat
        prev = text
        text = mw.ustring.gsub(text, "{{[^{}]-}}", "")
        count = count + 1
    until text == prev or count > 10
    -- 2. 使用安全占位符保護連結,剔除圖片
    repeat
        prev = text
        text = mw.ustring.gsub(text, "%[%[%s*([^%[%]]-)%s*%]%]", function(inner)
            local low = mw.ustring.lower(inner)
            if low:match("^file:") or low:match("^image:") or
              low:match("^文件:") or low:match("^圖像:") or
              low:match("^category:") or low:match("^分類:") then
                return ""
            end
            return "LINKSTART" .. inner .. "LINKEND"
        end)
    until text == prev
    -- 3. 清理剩餘雜質
    text = mw.ustring.gsub(text, "\n==+.-==+", " ")
    text = mw.ustring.gsub(text, "\n%s*[*#:]+", " ")
    text = mw.ustring.gsub(text, "\n+", " ")
    text = mw.ustring.gsub(text, "%s+", " ")
    -- 4. 初步恢復連結標籤
    text = mw.ustring.gsub(text, "LINKSTART", "[[")
    text = mw.ustring.gsub(text, "LINKEND", "]]")
    return mw.text.trim(text)
end


function p.getSummaryAndImage(frame)
function p.getSummaryAndImage(frame)
     local pageName = frame.args[1] or ""
     local pageName = frame.args[1] or ""
     local title = mw.title.new(pageName)
     local title = mw.title.new(pageName)
   
     if not title or not title.exists then return "" end
     if not title or not title.exists then return "页面不存在" end
 
     local content = title:getContent()
    -- 【調整】讀取前 1000 字以確保跨過開頭的表格/模板區
     local rawContent = title:getContent() or ""
    local limitedContent = mw.ustring.sub(rawContent, 1, 1000)


     -- 1. 提取第一张图(提取动作在清理前完成)
     -- 1. 提取第一張縮圖名
     local firstImage = string.match(content, "%[%[([Ff]ile:.-)[|%]]") or  
     local firstImage = mw.ustring.match(limitedContent, "%[%[%s*[Ff]ile%s*:([^|%]%s\n]+)") or  
                       string.match(content, "%[%[([Ii]mage:.-)[|%]]") or
                       mw.ustring.match(limitedContent, "%[%[%s*文件%s*:([^|%]%s\n]+)") or
                       string.match(content, "%[%[(文件:.-)[|%]]") or
                       mw.ustring.match(limitedContent, "%[%[%s*[Ii]mage%s*:([^|%]%s\n]+)")
                      string.match(content, "%[%[(图像:.-)[|%]]")


     -- 2. 处理内容:预加载 1000 字符
     -- 2. 執行清理(此時會得到過濾掉干擾後的純文字)
     local rawText = mw.ustring.sub(content, 1, 1000)
     local cleanText = cleanContent(limitedContent)
      
      
     -- A. 高级清理图片:处理带有长描述和嵌套链接的图片
     -- 3. 【調整】最終截取約 80 字
     -- 循环匹配 [[File: ... ]] 结构,尽量处理嵌套情况
     local targetLen = 80
    local function stripImages(text)
     local summary = ""
        -- 匹配 [[(File|文件|Image|图像): ... ]]
        -- 此正则通过排除法尽可能多地包含字符,直到匹配到对应的闭合括号
        local patterns = {"%[%[[Ff]ile:.-%]%]", "%[%[[Ii]mage:.-%]%]", "%[%[文件:.-%]%]", "%[%[图像:.-%]%]" }
        for _, pat in ipairs(patterns) do
            while string.find(text, pat) do
                text = string.gsub(text, pat, "")
            end
        end
        return text
    end
    rawText = stripImages(rawText)
   
    -- B. 移除 Wiki 列表符号 (*, #, :, ;)
    rawText = string.gsub(rawText, "\n%s*[*#:]+", " ")
   
    -- C. 移除标题符号 == 标题 ==
    rawText = string.gsub(rawText, "\n==+.-==+", " ")
   
    -- D. 移除模板 {{...}}
     local limit = 15
    while string.find(rawText, "{{.-}}") and limit > 0 do
        rawText = string.gsub(rawText, "{{[^{}]+}}", "")
        limit = limit - 1
    end
   
    -- E. 移除引用 <ref> 和 注释
    rawText = string.gsub(rawText, "<ref.-</ref>", "")
    rawText = string.gsub(rawText, "<ref.->", "")
    rawText = string.gsub(rawText, "<!%-%-.-%-%->", "")
   
    -- F. 清理多余的空白符、换行符和特殊占位符(如 &nbsp;)
    rawText = string.gsub(rawText, "&nbsp;", " ")
    rawText = string.gsub(rawText, "\n", " ")
    rawText = string.gsub(rawText, "%s+", " ")
      
      
    local cleanText = mw.text.trim(rawText)
    -- 3. 截断逻辑:150 字
    local targetLen = 150
    local summary = ""
     if mw.ustring.len(cleanText) <= targetLen then
     if mw.ustring.len(cleanText) <= targetLen then
         summary = cleanText
         summary = cleanText
     else
     else
         summary = mw.ustring.sub(cleanText, 1, targetLen)
        -- 截取前 80 字並加上省略號
         local restText = mw.ustring.sub(cleanText, targetLen + 1)
         summary = mw.ustring.sub(cleanText, 1, targetLen) .. "..."
       
        -- 閉合加粗標籤 '''
         local _, opens = mw.ustring.gsub(summary, "'''", "")
        if opens % 2 ~= 0 then summary = summary .. "'''" end
          
          
         -- 链接延展逻辑
         -- 閉合或清理截斷的連結 [[...
         local lastOpen = 0
         if mw.ustring.match(summary, "%[%[[^%]]*$") then
        local lastClose = 0
             summary = mw.ustring.gsub(summary, "%[%[[^%]]*$", "") .. "..."
        local tempPos = 0
        while true do
            local found = string.find(summary, "%[%[", tempPos + 1)
            if not found then break end
             lastOpen = found
            tempPos = found
        end
        tempPos = 0
        while true do
            local found = string.find(summary, "%]%]", tempPos + 1)
            if not found then break end
            lastClose = found
            tempPos = found
        end
 
        if lastOpen > lastClose then
            local endOfLink = string.find(restText, "%]%]")
            if endOfLink then
                summary = summary .. string.sub(restText, 1, endOfLink + 1)
            end
         end
         end
        summary = summary .. "..."
     end
     end


     -- 4. 渲染
     -- 4. 渲染 HTML
     local container = mw.html.create('div')
     local res = mw.html.create('div'):css({
        :css({
        ['display'] = 'flow-root',  
            ['margin'] = '15px 0',
        ['line-height'] = '1.5',
            ['background'] = 'transparent',
        ['margin-bottom'] = '1em' -- 也可以通过这里控制整个组件下方的间距
            ['display'] = 'flow-root'
    })
        })
   
    -- 标题
    res:tag('div')
      :css({['font-size'] = '1.15em', ['font-weight'] = 'bold', ['margin-bottom'] = '4px'})
      :wikitext('[[' .. pageName .. ']]')


     -- 原生格式图片
     -- 右侧缩略图
     if firstImage then
     if firstImage then
         container:wikitext('[[' .. firstImage .. '|150px|right|link=' .. pageName .. ']]')
         res:wikitext('[[File:' .. firstImage .. '|100px|right|link=' .. pageName .. ']]')
     end
     end


     -- 标题
     -- 摘要正文:在末尾增加空行
     container:tag('div')
     -- 方法:在 summary 字符串后直接加上 <br /> 或 \n\n
        :css({
     res:tag('div')
            ['font-size'] = '1.4em',
      :wikitext(summary .. '<br /><br />')  
            ['font-weight'] = 'bold',
            ['margin-bottom'] = '8px'
        })
        :wikitext('[[' .. pageName .. ']]')
 
    -- 摘要文本
     container:tag('div')
        :css({
            ['line-height'] = '1.6',
            ['color'] = '#202122'
        })
        :wikitext(summary)


     return tostring(container)
     return tostring(res)
end
end


return p
return p