Shortcode rewrite, take 2
This commit contains a restructuring and partial rewrite of the shortcode handling.
Prior to this commit rendering of the page content was mingled with handling of the shortcodes. This led to several oddities.
The new flow is:
1. Shortcodes are extracted from page and replaced with placeholders.
2. Shortcodes are processed and rendered
3. Page is processed
4. The placeholders are replaced with the rendered shortcodes
The handling of summaries is also made simpler by this.
This commit also introduces some other chenges:
1. distinction between shortcodes that need further processing and those who do not:
* `{{< >}}`: Typically raw HTML. Will not be processed.
* `{{% %}}`: Will be processed by the page's markup engine (Markdown or (infuture) Asciidoctor)
The above also involves a new shortcode-parser, with lexical scanning inspired by Rob Pike's talk called "Lexical Scanning in Go",
which should be easier to understand, give better error messages and perform better.
2. If you want to exclude a shortcode from being processed (for documentation etc.), the inner part of the shorcode must be commented out, i.e. `{{%/* movie 47238zzb */%}}`. See the updated shortcode section in the documentation for further examples.
The new parser supports nested shortcodes. This isn't new, but has two related design choices worth mentioning:
* The shortcodes will be rendered individually, so If both `{{< >}}` and `{{% %}}` are used in the nested hierarchy, one will be passed through the page's markdown processor, the other not.
* To avoid potential costly overhead of always looking far ahead for a possible closing tag, this implementation looks at the template itself, and is branded as a container with inner content if it contains a reference to `.Inner`
Fixes #565
Fixes #480
Fixes #461
And probably some others.
2014-10-27 16:48:30 -04:00
|
|
|
// Copyright © 2013-14 Steve Francia <spf@spf13.com>.
|
2013-12-05 09:42:29 -05:00
|
|
|
//
|
|
|
|
// Licensed under the Simple Public License, Version 2.0 (the "License");
|
|
|
|
// you may not use this file except in compliance with the License.
|
|
|
|
// You may obtain a copy of the License at
|
|
|
|
// http://opensource.org/licenses/Simple-2.0
|
|
|
|
//
|
|
|
|
// Unless required by applicable law or agreed to in writing, software
|
|
|
|
// distributed under the License is distributed on an "AS IS" BASIS,
|
|
|
|
// WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
|
|
|
|
// See the License for the specific language governing permissions and
|
|
|
|
// limitations under the License.
|
|
|
|
|
|
|
|
package helpers
|
|
|
|
|
|
|
|
import (
|
2014-01-29 17:50:31 -05:00
|
|
|
"bytes"
|
2015-03-27 12:05:17 -04:00
|
|
|
"crypto/sha1"
|
2014-05-07 12:38:14 -04:00
|
|
|
"fmt"
|
2015-04-15 14:31:05 -04:00
|
|
|
"github.com/spf13/hugo/hugofs"
|
|
|
|
jww "github.com/spf13/jwalterweatherman"
|
|
|
|
"github.com/spf13/viper"
|
2015-03-29 06:55:46 -04:00
|
|
|
"io"
|
2015-03-27 12:05:17 -04:00
|
|
|
"io/ioutil"
|
2014-01-29 17:50:31 -05:00
|
|
|
"os/exec"
|
2015-04-14 13:58:37 -04:00
|
|
|
"path/filepath"
|
2015-04-15 14:31:05 -04:00
|
|
|
"sort"
|
2014-01-29 17:50:31 -05:00
|
|
|
"strings"
|
2013-12-05 09:42:29 -05:00
|
|
|
)
|
|
|
|
|
Shortcode rewrite, take 2
This commit contains a restructuring and partial rewrite of the shortcode handling.
Prior to this commit rendering of the page content was mingled with handling of the shortcodes. This led to several oddities.
The new flow is:
1. Shortcodes are extracted from page and replaced with placeholders.
2. Shortcodes are processed and rendered
3. Page is processed
4. The placeholders are replaced with the rendered shortcodes
The handling of summaries is also made simpler by this.
This commit also introduces some other chenges:
1. distinction between shortcodes that need further processing and those who do not:
* `{{< >}}`: Typically raw HTML. Will not be processed.
* `{{% %}}`: Will be processed by the page's markup engine (Markdown or (infuture) Asciidoctor)
The above also involves a new shortcode-parser, with lexical scanning inspired by Rob Pike's talk called "Lexical Scanning in Go",
which should be easier to understand, give better error messages and perform better.
2. If you want to exclude a shortcode from being processed (for documentation etc.), the inner part of the shorcode must be commented out, i.e. `{{%/* movie 47238zzb */%}}`. See the updated shortcode section in the documentation for further examples.
The new parser supports nested shortcodes. This isn't new, but has two related design choices worth mentioning:
* The shortcodes will be rendered individually, so If both `{{< >}}` and `{{% %}}` are used in the nested hierarchy, one will be passed through the page's markdown processor, the other not.
* To avoid potential costly overhead of always looking far ahead for a possible closing tag, this implementation looks at the template itself, and is branded as a container with inner content if it contains a reference to `.Inner`
Fixes #565
Fixes #480
Fixes #461
And probably some others.
2014-10-27 16:48:30 -04:00
|
|
|
const pygmentsBin = "pygmentize"
|
2013-12-05 09:42:29 -05:00
|
|
|
|
2014-12-26 10:07:03 -05:00
|
|
|
// HasPygments checks to see if Pygments is installed and available
|
|
|
|
// on the system.
|
Shortcode rewrite, take 2
This commit contains a restructuring and partial rewrite of the shortcode handling.
Prior to this commit rendering of the page content was mingled with handling of the shortcodes. This led to several oddities.
The new flow is:
1. Shortcodes are extracted from page and replaced with placeholders.
2. Shortcodes are processed and rendered
3. Page is processed
4. The placeholders are replaced with the rendered shortcodes
The handling of summaries is also made simpler by this.
This commit also introduces some other chenges:
1. distinction between shortcodes that need further processing and those who do not:
* `{{< >}}`: Typically raw HTML. Will not be processed.
* `{{% %}}`: Will be processed by the page's markup engine (Markdown or (infuture) Asciidoctor)
The above also involves a new shortcode-parser, with lexical scanning inspired by Rob Pike's talk called "Lexical Scanning in Go",
which should be easier to understand, give better error messages and perform better.
2. If you want to exclude a shortcode from being processed (for documentation etc.), the inner part of the shorcode must be commented out, i.e. `{{%/* movie 47238zzb */%}}`. See the updated shortcode section in the documentation for further examples.
The new parser supports nested shortcodes. This isn't new, but has two related design choices worth mentioning:
* The shortcodes will be rendered individually, so If both `{{< >}}` and `{{% %}}` are used in the nested hierarchy, one will be passed through the page's markdown processor, the other not.
* To avoid potential costly overhead of always looking far ahead for a possible closing tag, this implementation looks at the template itself, and is branded as a container with inner content if it contains a reference to `.Inner`
Fixes #565
Fixes #480
Fixes #461
And probably some others.
2014-10-27 16:48:30 -04:00
|
|
|
func HasPygments() bool {
|
2014-01-29 17:50:31 -05:00
|
|
|
if _, err := exec.LookPath(pygmentsBin); err != nil {
|
Shortcode rewrite, take 2
This commit contains a restructuring and partial rewrite of the shortcode handling.
Prior to this commit rendering of the page content was mingled with handling of the shortcodes. This led to several oddities.
The new flow is:
1. Shortcodes are extracted from page and replaced with placeholders.
2. Shortcodes are processed and rendered
3. Page is processed
4. The placeholders are replaced with the rendered shortcodes
The handling of summaries is also made simpler by this.
This commit also introduces some other chenges:
1. distinction between shortcodes that need further processing and those who do not:
* `{{< >}}`: Typically raw HTML. Will not be processed.
* `{{% %}}`: Will be processed by the page's markup engine (Markdown or (infuture) Asciidoctor)
The above also involves a new shortcode-parser, with lexical scanning inspired by Rob Pike's talk called "Lexical Scanning in Go",
which should be easier to understand, give better error messages and perform better.
2. If you want to exclude a shortcode from being processed (for documentation etc.), the inner part of the shorcode must be commented out, i.e. `{{%/* movie 47238zzb */%}}`. See the updated shortcode section in the documentation for further examples.
The new parser supports nested shortcodes. This isn't new, but has two related design choices worth mentioning:
* The shortcodes will be rendered individually, so If both `{{< >}}` and `{{% %}}` are used in the nested hierarchy, one will be passed through the page's markdown processor, the other not.
* To avoid potential costly overhead of always looking far ahead for a possible closing tag, this implementation looks at the template itself, and is branded as a container with inner content if it contains a reference to `.Inner`
Fixes #565
Fixes #480
Fixes #461
And probably some others.
2014-10-27 16:48:30 -04:00
|
|
|
return false
|
|
|
|
}
|
|
|
|
return true
|
|
|
|
}
|
|
|
|
|
2014-12-26 10:07:03 -05:00
|
|
|
// Highlight takes some code and returns highlighted code.
|
2015-04-15 14:31:05 -04:00
|
|
|
func Highlight(code, lang, optsStr string) string {
|
2014-03-31 13:23:34 -04:00
|
|
|
|
Shortcode rewrite, take 2
This commit contains a restructuring and partial rewrite of the shortcode handling.
Prior to this commit rendering of the page content was mingled with handling of the shortcodes. This led to several oddities.
The new flow is:
1. Shortcodes are extracted from page and replaced with placeholders.
2. Shortcodes are processed and rendered
3. Page is processed
4. The placeholders are replaced with the rendered shortcodes
The handling of summaries is also made simpler by this.
This commit also introduces some other chenges:
1. distinction between shortcodes that need further processing and those who do not:
* `{{< >}}`: Typically raw HTML. Will not be processed.
* `{{% %}}`: Will be processed by the page's markup engine (Markdown or (infuture) Asciidoctor)
The above also involves a new shortcode-parser, with lexical scanning inspired by Rob Pike's talk called "Lexical Scanning in Go",
which should be easier to understand, give better error messages and perform better.
2. If you want to exclude a shortcode from being processed (for documentation etc.), the inner part of the shorcode must be commented out, i.e. `{{%/* movie 47238zzb */%}}`. See the updated shortcode section in the documentation for further examples.
The new parser supports nested shortcodes. This isn't new, but has two related design choices worth mentioning:
* The shortcodes will be rendered individually, so If both `{{< >}}` and `{{% %}}` are used in the nested hierarchy, one will be passed through the page's markdown processor, the other not.
* To avoid potential costly overhead of always looking far ahead for a possible closing tag, this implementation looks at the template itself, and is branded as a container with inner content if it contains a reference to `.Inner`
Fixes #565
Fixes #480
Fixes #461
And probably some others.
2014-10-27 16:48:30 -04:00
|
|
|
if !HasPygments() {
|
2014-04-09 15:08:47 -04:00
|
|
|
jww.WARN.Println("Highlighting requires Pygments to be installed and in the path")
|
2014-01-29 17:50:31 -05:00
|
|
|
return code
|
|
|
|
}
|
2013-12-05 09:42:29 -05:00
|
|
|
|
2015-04-15 14:31:05 -04:00
|
|
|
options, err := parsePygmentsOpts(optsStr)
|
2015-03-29 06:55:46 -04:00
|
|
|
|
2015-04-15 14:31:05 -04:00
|
|
|
if err != nil {
|
|
|
|
jww.ERROR.Print(err.Error())
|
|
|
|
return code
|
2015-03-29 06:55:46 -04:00
|
|
|
}
|
|
|
|
|
2015-03-27 12:05:17 -04:00
|
|
|
// Try to read from cache first
|
2015-03-29 06:55:46 -04:00
|
|
|
hash := sha1.New()
|
|
|
|
io.WriteString(hash, code)
|
2015-04-15 14:31:05 -04:00
|
|
|
io.WriteString(hash, lang)
|
|
|
|
io.WriteString(hash, options)
|
2015-03-29 06:55:46 -04:00
|
|
|
|
2015-04-15 14:31:05 -04:00
|
|
|
fs := hugofs.OsFs
|
|
|
|
|
2015-04-15 20:27:37 -04:00
|
|
|
cacheDir := viper.GetString("CacheDir")
|
|
|
|
var cachefile string
|
2015-03-27 12:05:17 -04:00
|
|
|
|
2015-04-15 20:27:37 -04:00
|
|
|
if cacheDir != "" {
|
|
|
|
cachefile = filepath.Join(cacheDir, fmt.Sprintf("pygments-%x", hash.Sum(nil)))
|
|
|
|
|
|
|
|
exists, err := Exists(cachefile, fs)
|
2015-03-27 12:05:17 -04:00
|
|
|
if err != nil {
|
|
|
|
jww.ERROR.Print(err.Error())
|
|
|
|
return code
|
|
|
|
}
|
2015-04-15 20:27:37 -04:00
|
|
|
if exists {
|
|
|
|
f, err := fs.Open(cachefile)
|
|
|
|
if err != nil {
|
|
|
|
jww.ERROR.Print(err.Error())
|
|
|
|
return code
|
|
|
|
}
|
|
|
|
|
|
|
|
s, err := ioutil.ReadAll(f)
|
|
|
|
if err != nil {
|
|
|
|
jww.ERROR.Print(err.Error())
|
|
|
|
return code
|
|
|
|
}
|
|
|
|
|
|
|
|
return string(s)
|
|
|
|
}
|
2015-03-27 12:05:17 -04:00
|
|
|
}
|
|
|
|
|
|
|
|
// No cache file, render and cache it
|
2014-01-29 17:50:31 -05:00
|
|
|
var out bytes.Buffer
|
|
|
|
var stderr bytes.Buffer
|
2014-05-07 12:38:14 -04:00
|
|
|
|
2015-07-03 17:51:43 -04:00
|
|
|
var langOpt string
|
|
|
|
if lang == "" {
|
|
|
|
langOpt = "-g" // Try guessing the language
|
|
|
|
} else {
|
2015-10-12 14:47:06 -04:00
|
|
|
langOpt = "-l" + lang
|
2015-07-03 17:51:43 -04:00
|
|
|
}
|
|
|
|
|
|
|
|
cmd := exec.Command(pygmentsBin, langOpt, "-fhtml", "-O", options)
|
2014-01-29 17:50:31 -05:00
|
|
|
cmd.Stdin = strings.NewReader(code)
|
|
|
|
cmd.Stdout = &out
|
|
|
|
cmd.Stderr = &stderr
|
2013-12-05 09:42:29 -05:00
|
|
|
|
2014-01-29 17:50:31 -05:00
|
|
|
if err := cmd.Run(); err != nil {
|
2014-03-31 13:23:34 -04:00
|
|
|
jww.ERROR.Print(stderr.String())
|
2014-01-29 17:50:31 -05:00
|
|
|
return code
|
|
|
|
}
|
2013-12-05 09:42:29 -05:00
|
|
|
|
2015-10-14 17:10:50 -04:00
|
|
|
str := out.String()
|
|
|
|
|
|
|
|
// inject code tag into Pygments output
|
|
|
|
if lang != "" && strings.Contains(str, "<pre>") {
|
|
|
|
codeTag := fmt.Sprintf(`<pre><code class="language-%s" data-lang="%s">`, lang, lang)
|
|
|
|
str = strings.Replace(str, "<pre>", codeTag, 1)
|
|
|
|
str = strings.Replace(str, "</pre>", "</code></pre>", 1)
|
|
|
|
}
|
|
|
|
|
2015-04-15 20:27:37 -04:00
|
|
|
if cachefile != "" {
|
|
|
|
// Write cache file
|
2015-10-14 17:10:50 -04:00
|
|
|
if err := WriteToDisk(cachefile, strings.NewReader(str), fs); err != nil {
|
2015-04-15 20:27:37 -04:00
|
|
|
jww.ERROR.Print(stderr.String())
|
|
|
|
}
|
2015-03-27 12:05:17 -04:00
|
|
|
}
|
|
|
|
|
2015-10-14 17:10:50 -04:00
|
|
|
return str
|
2013-12-05 09:42:29 -05:00
|
|
|
}
|
2015-04-15 14:31:05 -04:00
|
|
|
|
|
|
|
var pygmentsKeywords = make(map[string]bool)
|
|
|
|
|
|
|
|
func init() {
|
|
|
|
pygmentsKeywords["style"] = true
|
|
|
|
pygmentsKeywords["encoding"] = true
|
|
|
|
pygmentsKeywords["noclasses"] = true
|
|
|
|
pygmentsKeywords["hl_lines"] = true
|
|
|
|
pygmentsKeywords["linenos"] = true
|
2015-05-06 20:55:28 -04:00
|
|
|
pygmentsKeywords["classprefix"] = true
|
2015-04-15 14:31:05 -04:00
|
|
|
}
|
|
|
|
|
|
|
|
func parsePygmentsOpts(in string) (string, error) {
|
|
|
|
|
|
|
|
in = strings.Trim(in, " ")
|
|
|
|
|
|
|
|
style := viper.GetString("PygmentsStyle")
|
|
|
|
|
|
|
|
noclasses := "true"
|
|
|
|
if viper.GetBool("PygmentsUseClasses") {
|
|
|
|
noclasses = "false"
|
|
|
|
}
|
|
|
|
|
|
|
|
if len(in) == 0 {
|
|
|
|
return fmt.Sprintf("style=%s,noclasses=%s,encoding=utf8", style, noclasses), nil
|
|
|
|
}
|
|
|
|
|
|
|
|
options := make(map[string]string)
|
|
|
|
|
|
|
|
o := strings.Split(in, ",")
|
|
|
|
for _, v := range o {
|
|
|
|
keyVal := strings.Split(v, "=")
|
|
|
|
key := strings.ToLower(strings.Trim(keyVal[0], " "))
|
|
|
|
if len(keyVal) != 2 || !pygmentsKeywords[key] {
|
|
|
|
return "", fmt.Errorf("invalid Pygments option: %s", key)
|
|
|
|
}
|
|
|
|
options[key] = keyVal[1]
|
|
|
|
}
|
|
|
|
|
|
|
|
if _, ok := options["style"]; !ok {
|
|
|
|
options["style"] = style
|
|
|
|
}
|
|
|
|
|
|
|
|
if _, ok := options["noclasses"]; !ok {
|
|
|
|
options["noclasses"] = noclasses
|
|
|
|
}
|
|
|
|
|
|
|
|
if _, ok := options["encoding"]; !ok {
|
|
|
|
options["encoding"] = "utf8"
|
|
|
|
}
|
|
|
|
|
|
|
|
var keys []string
|
|
|
|
for k := range options {
|
|
|
|
keys = append(keys, k)
|
|
|
|
}
|
|
|
|
sort.Strings(keys)
|
|
|
|
|
|
|
|
var optionsStr string
|
|
|
|
for i, k := range keys {
|
|
|
|
optionsStr += fmt.Sprintf("%s=%s", k, options[k])
|
|
|
|
if i < len(options)-1 {
|
|
|
|
optionsStr += ","
|
|
|
|
}
|
|
|
|
}
|
|
|
|
return optionsStr, nil
|
|
|
|
}
|