In [None]:
# Overview
# pandas has an options API configure and customize global behavior related to DataFrame display, data behavior and more.

# Options have a full “dotted-style”, case-insensitive name (e.g. display.max_rows). You can get/set options directly as attributes of the top-level options attribute:

In [1]:
import pandas as pd

pd.options.display.max_rows
pd.options.display.max_rows = 999

60

In [3]:
# The API is composed of 5 relevant functions, available directly from the pandas namespace:

In [None]:
# get_option() / set_option() - get/set the value of a single option.

# reset_option() - reset one or more options to their default value.

# describe_option() - print the descriptions of one or more options.

# option_context() - execute a codeblock with a set of options that revert to prior settings after execution.

In [None]:
# All of the functions above accept a regexp pattern (re.search style) as an argument, to match an unambiguous substring:

In [4]:
pd.get_option("display.chop_threshold")

In [5]:
pd.set_option("display.chop_threshold", 2)

In [6]:
pd.get_option("display.chop_threshold")

2

In [None]:
# The following will not work because it matches multiple option names, e.g. display.max_colwidth, display.max_rows, display.max_columns:

In [10]:
pd.get_option("max")

OptionError: Pattern matched multiple keys

In [None]:
# Available options

In [None]:
# You can get a list of available options and their descriptions with describe_option(). When called with no argument describe_option() will print out the descriptions for all available options.

In [7]:
pd.describe_option()

compute.use_bottleneck : bool
    Use the bottleneck library to accelerate if it is installed,
    the default is True
    Valid values: False,True
    [default: True] [currently: True]
compute.use_numba : bool
    Use the numba engine option for select operations if it is installed,
    the default is False
    Valid values: False,True
    [default: False] [currently: False]
compute.use_numexpr : bool
    Use the numexpr library to accelerate computation if it is installed,
    the default is True
    Valid values: False,True
    [default: True] [currently: True]
display.chop_threshold : float or None
    if set to a float value, all float values smaller than the given threshold
    will be displayed as exactly 0 by repr and friends.
    [default: None] [currently: 2]
display.colheader_justify : 'left'/'right'
    Controls the justification of column headers. used by DataFrameFormatter.
    [default: right] [currently: right]
display.date_dayfirst : boolean
    When True, prints and p

In [None]:
# Getting and setting options

In [None]:
# As described above, get_option() and set_option() are available from the pandas namespace. To change an option, call set_option('option regex', new_value).

In [8]:
pd.get_option("mode.sim_interactive")

False

In [9]:
pd.set_option("mode.sim_interactive", True)

In [10]:
pd.get_option("mode.sim_interactive")

True

In [None]:
# The option 'mode.sim_interactive' is mostly used for debugging purposes. You can use reset_option() to revert to a setting’s default value

In [12]:
pd.get_option("display.max_rows")
pd.set_option("display.max_rows", 999)

In [None]:
# It’s also possible to reset multiple options at once (using a regex):

In [13]:
pd.reset_option("^display")

In [None]:
# option_context() context manager has been exposed through the top-level API, allowing you to execute code with given option values. Option values are restored automatically when you exit the with block:

In [14]:
with pd.option_context("display.max_rows", 10, "display.max_columns", 5):
    print(pd.get_option("display.max_rows"))
    print(pd.get_option("display.max_columns"))

10
5


In [15]:
print(pd.get_option("display.max_rows"))

60


In [16]:
print(pd.get_option("display.max_columns"))

20


In [None]:
# Setting startup options in Python/IPython environment

In [None]:
# Using startup scripts for the Python/IPython environment to import pandas and set options makes working with pandas more efficient. To do this, create a .py or .ipy script in the startup directory of the desired profile. An example where the startup folder is in a default IPython profile can be found at:

In [None]:
# More information can be found in the IPython documentation. An example startup script for pandas is displayed below:

In [17]:
import pandas as pd

pd.set_option("display.max_rows", 999)
pd.set_option("display.precision", 5)

In [None]:
# Frequently used options
# The following is a demonstrates the more frequently used display options.

# display.max_rows and display.max_columns sets the maximum number of rows and columns displayed when a frame is pretty-printed. Truncated lines are replaced by an ellipsis.

In [20]:
import numpy as np
df = pd.DataFrame(np.random.randn(7, 2))
pd.set_option("display.max_rows", 7)

In [None]:
# Once the display.max_rows is exceeded, the display.min_rows options determines how many rows are shown in the truncated repr.

In [21]:
pd.set_option("display.max_rows", 8)

pd.set_option("display.min_rows", 4)

In [22]:
df = pd.DataFrame(np.random.randn(7, 2))

df

Unnamed: 0,0,1
0,0.92976,0.40079
1,-0.69682,0.41961
2,0.36914,1.25783
3,0.77244,-0.85888
4,0.65936,1.09444
5,0.35393,-0.26496
6,1.13218,1.34001


In [23]:
df = pd.DataFrame(np.random.randn(9, 2))

In [24]:
pd.reset_option("display.max_rows")

In [None]:
# display.expand_frame_repr allows for the representation of a DataFrame to stretch across pages, wrapped over the all the columns.

In [25]:
df = pd.DataFrame(np.random.randn(5, 10))

pd.set_option("expand_frame_repr", True)

df

Unnamed: 0,0,1,2,3,4,5,6,7,8,9
0,-0.51942,-0.56897,-0.38225,-1.19588,-1.35953,-0.23347,0.57621,-1.26817,0.73253,-1.83991
1,-2.28512,-0.36524,0.3487,1.21298,0.86697,-0.3699,0.79202,1.00392,1.73456,-0.00927
2,-0.03039,1.36917,-0.41194,-2.19088,-1.62512,-2.00065,1.0094,-1.04161,1.01556,-0.36352
3,1.05493,0.04975,-0.17097,-0.55087,-0.94629,1.37259,0.13717,1.04655,0.80586,0.87431
4,-0.55741,1.43989,0.7969,0.0323,-0.20216,0.27017,0.05311,-0.16594,1.22618,-1.14937


In [26]:
pd.set_option("expand_frame_repr", False)

In [None]:
# display.large_repr displays a DataFrame that exceed max_columns or max_rows as a truncated frame or summary.

In [28]:
df = pd.DataFrame(np.random.randn(10, 10))

pd.set_option("display.max_rows", 5)
pd.set_option("large_repr", "truncate")

In [29]:
pd.set_option("large_repr", "info")

In [30]:
df

In [31]:
pd.reset_option("large_repr")

In [32]:

pd.reset_option("display.max_rows")

In [None]:
# display.max_colwidth sets the maximum width of columns. Cells of this length or longer will be truncated with an ellipsis.

In [33]:
df = pd.DataFrame(
    np.array(
        [
            ["foo", "bar", "bim", "uncomfortably long string"],
            ["horse", "cow", "banana", "apple"],
        ]
    )
)

In [34]:
pd.set_option("max_colwidth", 40)

In [None]:
# display.max_info_columns sets a threshold for the number of columns displayed when calling info().

In [35]:
df = pd.DataFrame(np.random.randn(10, 10))

pd.set_option("max_info_columns", 11)

In [None]:
# Number formatting
# pandas also allows you to set how numbers are displayed in the console. This option is not set through the set_options API.

# Use the set_eng_float_format function to alter the floating-point formatting of pandas objects to produce a particular format.

In [2]:
import numpy as np
import pandas as pd

pd.set_eng_float_format(accuracy=3, use_eng_prefix=True)

s = pd.Series(np.random.randn(5), index=["a", "b", "c", "d", "e"])

s / 1.0e3

a    -19.529u
b    735.098u
c     50.370u
d     19.013u
e   -955.444u
dtype: float64

In [None]:
# Unicode formatting

In [None]:
# Some East Asian countries use Unicode characters whose width corresponds to two Latin characters. If a DataFrame or Series contains these characters, the default output mode may not align them properly.

In [4]:
df = pd.DataFrame({"国籍": ["UK", "日本"], "名前": ["Alice", "しのぶ"]})


In [None]:
# Enabling display.unicode.east_asian_width allows pandas to check each character’s “East Asian Width” property. These characters can be aligned properly by setting this option to True. However, this will result in longer render times than the standard len function.

In [6]:
pd.set_option("display.unicode.east_asian_width", True)
df

Unnamed: 0,国籍,名前
0,UK,Alice
1,日本,しのぶ


In [None]:
# In addition, Unicode characters whose width is “ambiguous” can either be 1 or 2 characters wide depending on the terminal setting or encoding. The option display.unicode.ambiguous_as_wide can be used to handle the ambiguity.

# By default, an “ambiguous” character’s width, such as “¡” (inverted exclamation) in the example below, is taken to be 1.

In [7]:
df = pd.DataFrame({"a": ["xxx", "¡¡"], "b": ["yyy", "¡¡"]})

df

Unnamed: 0,a,b
0,xxx,yyy
1,¡¡,¡¡


In [None]:
# Enabling display.unicode.ambiguous_as_wide makes pandas interpret these characters’ widths to be 2. (Note that this option will only be effective when display.unicode.east_asian_width is enabled.)

# However, setting this option incorrectly for your terminal will cause these characters to be aligned incorrectly:

In [9]:
pd.set_option("display.unicode.ambiguous_as_wide", True)
df

Unnamed: 0,a,b
0,xxx,yyy
1,¡¡,¡¡


In [None]:
# Table schema display

In [37]:
# DataFrame and Series will publish a Table Schema representation by default. This can be enabled globally with the display.html.table_schema option:

pd.set_option("display.html.table_schema", True)